{"meta":{"query_hash":"0ee2f2a50fc7","filters":{"topic":"Biomedical Text Mining and Ontologies"},"cohort_total":1581,"direct_labels_cover":8,"predictions_cover":1581,"exported":1581,"export_cap":100000,"truncated":false,"label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12"},"permalink":"https://metacan.xera.ac/q/0ee2f2a50fc7","api":"https://metacan.xera.ac/api/v1/cohort?topic=Biomedical+Text+Mining+and+Ontologies"},"results":[{"id":"W103049419","doi":"10.1007/978-3-642-38288-8_14","title":"Bio2RDF Release 2: Improved Coverage, Interoperability and Provenance of Life Science Linked Data","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":131,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"SPARQL; Computer science; RDF; Named graph; Scripting language; Interoperability; Information retrieval; Ontology; Data integration; Linked data; Database; World Wide Web; Semantic Web","score_opus":0.023735725000197997,"score_gpt":0.26728766036294604,"score_spread":0.24355193536274805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W103049419","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027502963,0.003692122,0.4837391,0.007249542,0.00345152,0.001349227,0.17649262,0.27855384,0.017969063],"genre_scores_gemma":[0.043263085,0.0017517802,0.36958802,0.0017364693,0.00060009974,0.001023449,0.4985991,0.064607896,0.018830057],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99055415,0.0015722228,0.00106257,0.0012230116,0.0050348756,0.00055328594],"domain_scores_gemma":[0.964906,0.0115797855,0.0010245858,0.01374642,0.0078988625,0.000844507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023043977,0.0016840092,0.0015162578,0.0064925184,0.0011937472,0.0064721275,0.0039558695,0.0031887526,0.01178236],"category_scores_gemma":[0.053519975,0.0017826977,0.0018860426,0.0055137775,0.00095480226,0.007344995,0.003975569,0.0032940235,0.010851089],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001142339,0.00036662165,0.004445524,0.0010922543,0.00021004726,0.0005694972,0.0009664432,0.005394494,0.01738689,0.013740141,0.6802571,0.27442864],"study_design_scores_gemma":[0.00056226394,0.00017808564,0.005112124,0.00060942787,0.00017860354,0.0011296598,0.00019522486,0.02485976,0.046619456,0.015689848,0.9045838,0.00028185206],"about_ca_topic_score_codex":0.017003624,"about_ca_topic_score_gemma":0.01332286,"teacher_disagreement_score":0.023043977,"about_ca_system_score_codex":0.0013938693,"about_ca_system_score_gemma":0.0040951967,"threshold_uncertainty_score":0.12186962},"labels":[],"label_agreement":null},{"id":"W112165001","doi":"10.3233/978-1-58603-979-0-513","title":"Normalization of Reported Lab Results","year":2009,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Langley Environmental Partners Society","funders":"","keywords":"Normalization (sociology); Popularity; Computer science; Visualization; Cluster (spacecraft); Data science; Artificial intelligence; Medical physics; Data mining; Medicine; Psychology","score_opus":0.04135943977708643,"score_gpt":0.36997858517857934,"score_spread":0.3286191454014929,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W112165001","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2367612,0.010686587,0.48083165,0.005070333,0.010841853,0.008866093,0.1375118,0.03113907,0.07829143],"genre_scores_gemma":[0.33893964,0.0052595134,0.52786034,0.0012495809,0.0016560834,0.0081911115,0.096360795,0.003852282,0.016630663],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.96566784,0.0070524346,0.008065555,0.0049872273,0.013213156,0.0010137222],"domain_scores_gemma":[0.8879898,0.022544973,0.012217744,0.023465553,0.05303502,0.00074686215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02060055,0.002260442,0.0020184442,0.011006099,0.000949351,0.0056475904,0.0024372926,0.0008465492,0.010355278],"category_scores_gemma":[0.13112508,0.0006457905,0.002140582,0.013322171,0.001147594,0.0027258925,0.0020365461,0.0020146642,0.006370195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020885307,0.0006564218,0.14261197,0.0051299026,0.00097216654,0.0008920766,0.002783716,0.006081783,0.024339976,0.0116889505,0.07478868,0.7279658],"study_design_scores_gemma":[0.00028433025,0.0010339293,0.4156991,0.0016974534,0.0011910569,0.0038759427,0.0039604935,0.025397286,0.08263366,0.025776265,0.43788365,0.0005667777],"about_ca_topic_score_codex":0.0043792087,"about_ca_topic_score_gemma":0.002844408,"teacher_disagreement_score":0.02060055,"about_ca_system_score_codex":0.0024046244,"about_ca_system_score_gemma":0.0027246394,"threshold_uncertainty_score":0.1089474},"labels":[],"label_agreement":null},{"id":"W113870509","doi":"10.1055/s-0038-1634293","title":"Analysis of the Process of Encoding Guidelines: A Comparison of GLIF2 and GLIF3","year":2002,"lang":"en","type":"article","venue":"Methods of Information in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Medical Research Council; U.S. National Library of Medicine; Medical Research Council Canada","keywords":"Formality; Encoding (memory); Guideline; Ambiguity; Clinical Practice; Computer science; Process (computing); Knowledge translation; Natural language processing; Information retrieval; Medicine; Knowledge management; Artificial intelligence; Linguistics; Family medicine; Pathology; Programming language","score_opus":0.07976690041168044,"score_gpt":0.45291873516480835,"score_spread":0.3731518347531279,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W113870509","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7553509,0.00040655365,0.22227788,0.0017726002,0.000089249406,0.0017020566,0.0035388183,0.0043788576,0.010483113],"genre_scores_gemma":[0.67357844,0.0002418126,0.31739458,0.00031561806,0.000026863861,0.0010559553,0.0047641117,0.000746965,0.0018756379],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98945844,0.0064895954,0.00100351,0.0006914241,0.0020602334,0.00029668052],"domain_scores_gemma":[0.8496229,0.11233842,0.0101115545,0.0109705,0.016054286,0.00090239453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020469865,0.0005068441,0.00050682476,0.0034584766,0.0007512903,0.003220826,0.0010508497,0.0009994355,0.0020892026],"category_scores_gemma":[0.11665268,0.00038943582,0.00094276556,0.003452777,0.0015768304,0.004684436,0.0020930525,0.0014659152,0.00051126693],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029973362,0.00053511164,0.069452114,0.0027720802,0.00036378202,0.000961708,0.09619012,0.007993559,0.020502653,0.012931802,0.007484956,0.77781487],"study_design_scores_gemma":[0.001641906,0.0038950944,0.30956915,0.0031972926,0.0010061726,0.005420618,0.10792651,0.25223473,0.13180918,0.04390756,0.13794057,0.0014511718],"about_ca_topic_score_codex":0.0085285995,"about_ca_topic_score_gemma":0.008020848,"teacher_disagreement_score":0.020469865,"about_ca_system_score_codex":0.0035759404,"about_ca_system_score_gemma":0.0036980584,"threshold_uncertainty_score":0.10825628},"labels":[],"label_agreement":null},{"id":"W1270312698","doi":"10.1007/s13721-015-0090-5","title":"Mining clinical text for stroke prediction","year":2015,"lang":"en","type":"article","venue":"Network Modeling Analysis in Health Informatics and Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Island Health; University of Victoria","funders":"","keywords":"Triage; Computer science; Health informatics; Support vector machine; Stroke (engine); Artificial intelligence; Machine learning; Natural language processing; Data science; Data mining; Medicine; Medical emergency; Public health; Pathology","score_opus":0.0835249501411915,"score_gpt":0.3615183872883345,"score_spread":0.27799343714714303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1270312698","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2906791,0.010932266,0.43965474,0.010872429,0.0011848338,0.0019553918,0.21843627,0.01005471,0.01623029],"genre_scores_gemma":[0.642013,0.0045754127,0.21817568,0.0010525697,0.0005870911,0.0010035428,0.1284842,0.00034706524,0.003761382],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990947,0.00019287091,0.00019769512,0.00024187233,0.00021027624,0.000062502535],"domain_scores_gemma":[0.994624,0.0037825073,0.00054254924,0.00028054856,0.00058403343,0.00018633499],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010561569,0.000979612,0.00062780763,0.0070422795,0.00067754736,0.0011861343,0.00076101173,0.00096515386,0.0044287415],"category_scores_gemma":[0.009097753,0.00023225928,0.0011765368,0.0047497507,0.00026063857,0.0019839108,0.000858265,0.0009166394,0.0017013897],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009790225,0.00083462533,0.15884139,0.0033856996,0.0008540237,0.0051308507,0.00073706167,0.02882619,0.015194658,0.01502825,0.0981562,0.672032],"study_design_scores_gemma":[0.0002566977,0.00039850667,0.07404367,0.0020316339,0.0026827147,0.006795581,0.0014353192,0.5996394,0.026268616,0.14220496,0.14411396,0.0001289116],"about_ca_topic_score_codex":0.0035238054,"about_ca_topic_score_gemma":0.005628474,"teacher_disagreement_score":0.0070422795,"about_ca_system_score_codex":0.0006775416,"about_ca_system_score_gemma":0.0017831951,"threshold_uncertainty_score":0.0148156285},"labels":[],"label_agreement":null},{"id":"W131864321","doi":"","title":"Representation of disorders of the newborn infant by SNOMED CT.","year":2008,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"SickKids Foundation; University of Toronto","funders":"","keywords":"SNOMED CT; Terminology; Representation (politics); Artificial intelligence; Computer science; Systematized Nomenclature of Medicine; Natural language processing; Set (abstract data type); Knowledge representation and reasoning; Medicine; Linguistics; Programming language","score_opus":0.015976034824081973,"score_gpt":0.22983014340179506,"score_spread":0.2138541085777131,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W131864321","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048962343,0.024449736,0.57122034,0.006374668,0.0010069625,0.004259399,0.24971311,0.01541876,0.07859469],"genre_scores_gemma":[0.11810213,0.010122305,0.6951077,0.0027422714,0.00020722768,0.001561294,0.16465339,0.0014205823,0.0060830605],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998086,0.00063006947,0.00048684562,0.00017732497,0.00054536137,0.00007436565],"domain_scores_gemma":[0.995074,0.002741131,0.00055106153,0.0005260969,0.0009543816,0.00015327707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002674466,0.0010616452,0.000624938,0.01255321,0.0007879719,0.0021109038,0.0013469433,0.0012683378,0.0062460015],"category_scores_gemma":[0.009401774,0.00035650388,0.0016120402,0.006802417,0.00076096493,0.0023203793,0.0018897356,0.0012817432,0.0018756943],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093747315,0.00018024871,0.01875753,0.016214915,0.00081259303,0.0112813655,0.0086840205,0.012545483,0.04598517,0.120181575,0.24464981,0.5197698],"study_design_scores_gemma":[0.00009200732,0.00017492038,0.014668674,0.0051441705,0.0007093595,0.014212629,0.002173062,0.018673038,0.014166334,0.043174703,0.8865431,0.0002680078],"about_ca_topic_score_codex":0.012263747,"about_ca_topic_score_gemma":0.0204675,"teacher_disagreement_score":0.01255321,"about_ca_system_score_codex":0.001257819,"about_ca_system_score_gemma":0.0044586267,"threshold_uncertainty_score":0.024384737},"labels":[],"label_agreement":null},{"id":"W1435051560","doi":"10.1017/cbo9780511791277.004","title":"The how-to of Bayesian inference","year":2005,"lang":"en","type":"book-chapter","venue":"Cambridge University Press eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Bayesian inference; Inference; Bayesian probability; Artificial intelligence; Computer science","score_opus":0.018332187764320645,"score_gpt":0.2223079926521309,"score_spread":0.20397580488781025,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1435051560","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006222865,0.030619282,0.8669295,0.039466627,0.0033167936,0.00009721953,0.00046886184,0.0006616036,0.057817835],"genre_scores_gemma":[0.057918493,0.050276384,0.815051,0.025618762,0.010513437,0.00087649707,0.00080652,0.0015354217,0.037403397],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9874035,0.007743033,0.0005762451,0.001220446,0.0028301454,0.0002265897],"domain_scores_gemma":[0.9761405,0.019531848,0.00031145674,0.0024508217,0.0013008175,0.00026450766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015384673,0.0015001198,0.0015913994,0.0027594222,0.0018638133,0.00765683,0.0029779505,0.004057202,0.013294951],"category_scores_gemma":[0.045750476,0.001271669,0.0017140915,0.0026571392,0.014329142,0.014263847,0.0042115245,0.011509235,0.0077345166],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000098984,0.000009970194,0.000119872304,0.00024447206,0.000036907863,0.00005115429,0.00031903057,0.0013160037,0.000082243445,0.9158233,0.019931702,0.062055502],"study_design_scores_gemma":[0.000004149089,0.0000022201866,0.00003521755,0.000115579096,0.000005470189,0.00004260462,0.000026819118,0.0013549239,0.000071505776,0.94506186,0.053268354,0.0000113579],"about_ca_topic_score_codex":0.0033459587,"about_ca_topic_score_gemma":0.0024826154,"teacher_disagreement_score":0.015384673,"about_ca_system_score_codex":0.0029969208,"about_ca_system_score_gemma":0.0027354755,"threshold_uncertainty_score":0.0813629},"labels":[],"label_agreement":null},{"id":"W144859847","doi":"10.5555/1999416.1999475","title":"On simulating episodic events against a background of noise-like non-episodic events","year":2010,"lang":"en","type":"article","venue":"Summer Computer Simulation Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Noise (video); Event (particle physics); Computer science; Process (computing); Field (mathematics); Seismology; Physics; Geology; Artificial intelligence; Mathematics; Quantum mechanics","score_opus":0.03577969889429551,"score_gpt":0.3098867298111211,"score_spread":0.2741070309168256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W144859847","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28745297,0.0002860916,0.7026969,0.0010292616,0.000103421415,0.0002168324,0.00044190124,0.00069883873,0.007073886],"genre_scores_gemma":[0.87841374,0.0005203939,0.117391914,0.00026749063,0.000068704285,0.00024680115,0.0004810182,0.00013228341,0.0024777101],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99927527,0.0003410314,0.000043114123,0.00012637816,0.00013104781,0.00008320751],"domain_scores_gemma":[0.9900842,0.008580999,0.00051133067,0.0003196981,0.0003359478,0.00016792594],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024570543,0.0008478061,0.0008231037,0.00086900924,0.00060545997,0.0012064041,0.0013643767,0.001649996,0.001604161],"category_scores_gemma":[0.01346234,0.0004493285,0.00080296077,0.00096998113,0.0017282292,0.0023160602,0.0012468125,0.0012688725,0.00019477673],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003600432,0.000024855835,0.001029248,0.000020963993,0.000016090491,0.000027308195,0.00006492228,0.9900355,0.0002905763,0.005194481,0.00008832032,0.0031717469],"study_design_scores_gemma":[0.000009425061,0.000018353248,0.00013132137,0.0000042664296,0.0000052441105,0.00000716461,0.00001462996,0.99539256,0.00018807908,0.0040401053,0.0001842797,0.000004544702],"about_ca_topic_score_codex":0.018443674,"about_ca_topic_score_gemma":0.011527449,"teacher_disagreement_score":0.018443674,"about_ca_system_score_codex":0.0012645309,"about_ca_system_score_gemma":0.0012421907,"threshold_uncertainty_score":0.03667265},"labels":[],"label_agreement":null},{"id":"W1458921027","doi":"","title":"Ontology-Based Extraction and Summarization of Protein Mutation Impact Information","year":2010,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Automatic summarization; Computer science; Ontology; Information retrieval; Information extraction; Mutation; Pace; Data mining; Geography; Biology; Genetics","score_opus":0.006578853222055952,"score_gpt":0.27714733627980903,"score_spread":0.2705684830577531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1458921027","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.124486424,0.002692535,0.7964002,0.002377236,0.00044869658,0.0012181416,0.051686566,0.011284981,0.009405237],"genre_scores_gemma":[0.19235627,0.002622599,0.7355883,0.00024192031,0.00025578734,0.0005384114,0.06499907,0.00063277315,0.0027648122],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982942,0.00021968385,0.0004274254,0.00030426524,0.00065981847,0.000094699906],"domain_scores_gemma":[0.9960394,0.0013947418,0.00074695674,0.00048376905,0.001206969,0.00012816038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017948544,0.0009014318,0.0010169661,0.015455084,0.00076970836,0.0018885075,0.00089680386,0.0006517634,0.0014422587],"category_scores_gemma":[0.008268399,0.0003222896,0.001285783,0.009566513,0.00035871167,0.0030411382,0.0012635654,0.0010051085,0.0008836957],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040458073,0.00040109054,0.01317062,0.0024059454,0.00048700953,0.0014898172,0.0024334576,0.009942805,0.082320556,0.016698316,0.033124886,0.83712095],"study_design_scores_gemma":[0.00025622454,0.00056880387,0.06573948,0.001126444,0.002958806,0.0033804716,0.0052224738,0.2797298,0.17097068,0.12304273,0.34647223,0.0005319024],"about_ca_topic_score_codex":0.0042071613,"about_ca_topic_score_gemma":0.0053341785,"teacher_disagreement_score":0.015455084,"about_ca_system_score_codex":0.0009955373,"about_ca_system_score_gemma":0.0023982134,"threshold_uncertainty_score":0.0094922185},"labels":[],"label_agreement":null},{"id":"W1480287196","doi":"10.1186/1471-2105-7-356","title":"New directions in biomedical text annotation: definitions, guidelines and corpus construction","year":2006,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":169,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"U.S. National Library of Medicine; Natural Sciences and Engineering Research Council of Canada; National Institutes of Health; National Science Foundation","keywords":"Annotation; Computer science; Task (project management); Information retrieval; Natural language processing; Categorization; Set (abstract data type); Biomedical text mining; Executable; Focus (optics); Artificial intelligence; Unified Medical Language System; Text corpus; Text mining","score_opus":0.03638827609910656,"score_gpt":0.27885588333730144,"score_spread":0.24246760723819488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1480287196","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006766012,0.009523025,0.92201364,0.049178652,0.0011525361,0.0026846072,0.0014204694,0.0017181755,0.005542841],"genre_scores_gemma":[0.011935683,0.0022064364,0.97548467,0.0019562116,0.0005332645,0.0046173115,0.0018220157,0.00047744188,0.0009670536],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.789187,0.14583577,0.03294713,0.012357975,0.018236658,0.0014355531],"domain_scores_gemma":[0.51539034,0.2991873,0.027742898,0.06283832,0.08872668,0.006114454],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.21495683,0.0024292807,0.0028096174,0.023429336,0.005309296,0.01766306,0.009875565,0.0068940474,0.0035025233],"category_scores_gemma":[0.293678,0.0030351805,0.0018450463,0.017760256,0.0293193,0.03766417,0.01180518,0.012632048,0.0035989748],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003403772,0.00047734013,0.007899763,0.012313337,0.00012411176,0.0009237124,0.0500807,0.0038545392,0.015610263,0.3488904,0.06746679,0.4920187],"study_design_scores_gemma":[0.00016918557,0.00027294687,0.005266161,0.011287324,0.0001398065,0.0013059698,0.019652218,0.022555137,0.013032159,0.50467914,0.42114982,0.0004901489],"about_ca_topic_score_codex":0.0073177093,"about_ca_topic_score_gemma":0.009847852,"teacher_disagreement_score":0.21495683,"about_ca_system_score_codex":0.0073184646,"about_ca_system_score_gemma":0.017111884,"threshold_uncertainty_score":0.968098},"labels":[],"label_agreement":null},{"id":"W1481389199","doi":"10.1136/amiajnl-2013-001636","title":"Literature review of SNOMED CT use","year":2013,"lang":"en","type":"review","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":177,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Systematized Nomenclature of Medicine; Implementation; Computer science; Medicine; MEDLINE; Domain (mathematical analysis); Information retrieval; Medical physics; Terminology; Software engineering","score_opus":0.0225821128250893,"score_gpt":0.33241396151467795,"score_spread":0.30983184868958863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1481389199","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008307685,0.9956045,0.0003166013,0.00096093846,0.00032435943,0.00006869227,0.00056123064,0.000014732563,0.0013181497],"genre_scores_gemma":[0.008646335,0.9869653,0.0012920036,0.001476247,0.00034675002,0.00015361793,0.0008665623,0.000021790858,0.00023138241],"study_design_codex":"systematic_review","study_design_gemma":"not_applicable","domain_scores_codex":[0.98284274,0.0055968757,0.006520959,0.0013243583,0.0034182013,0.00029681483],"domain_scores_gemma":[0.8932551,0.08411138,0.011535103,0.0015710475,0.008852624,0.000674729],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0144766215,0.0014159739,0.0030164805,0.043042388,0.0008791152,0.0033727556,0.0029349767,0.0021733742,0.009966039],"category_scores_gemma":[0.07022895,0.0007606728,0.0035711383,0.03388907,0.0017655132,0.004666637,0.0024304122,0.0012957582,0.0014243401],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017420869,0.000027374914,0.0017073489,0.7016974,0.001993346,0.0005630561,0.0014637938,0.0002035508,0.0005059084,0.0019595705,0.01584483,0.27385968],"study_design_scores_gemma":[0.000041173513,0.00010779228,0.0058109355,0.7809905,0.006859155,0.001975966,0.0013979982,0.00008889663,0.0005262324,0.0013068767,0.20084184,0.00005247102],"about_ca_topic_score_codex":0.004768458,"about_ca_topic_score_gemma":0.009882319,"teacher_disagreement_score":0.043042388,"about_ca_system_score_codex":0.0038863812,"about_ca_system_score_gemma":0.010716411,"threshold_uncertainty_score":0.07656062},"labels":[],"label_agreement":null},{"id":"W1481436478","doi":"10.1007/978-3-642-13059-5_10","title":"MeSH Represented MEDLINE Query Results","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Information retrieval; Terminology; Usability; Unified Medical Language System; Controlled vocabulary; MEDLINE; Domain (mathematical analysis); World Wide Web; Human–computer interaction","score_opus":0.01643965484155776,"score_gpt":0.2721156647028805,"score_spread":0.2556760098613228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1481436478","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01974099,0.0050368803,0.018992882,0.0015110668,0.0017375993,0.0008429891,0.9058339,0.010686327,0.035617404],"genre_scores_gemma":[0.111421235,0.0053014234,0.06900504,0.00090666907,0.00035941848,0.0012319097,0.7838375,0.002034749,0.02590199],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99861395,0.00019213204,0.00039884748,0.0003228321,0.00033639118,0.00013581237],"domain_scores_gemma":[0.998133,0.000948701,0.00014842451,0.00011627719,0.0005707556,0.00008294215],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.00069804373,0.0019654462,0.0011738462,0.01792766,0.00091897085,0.0030154078,0.0011621342,0.0017130752,0.067919075],"category_scores_gemma":[0.007987692,0.0003888736,0.0012850343,0.01282661,0.0005580548,0.0023111799,0.001747805,0.00077510695,0.023310628],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004261466,0.0003516424,0.006254172,0.024387693,0.0006245509,0.0047820536,0.0016537252,0.008085584,0.03758559,0.045626428,0.61448807,0.251899],"study_design_scores_gemma":[0.00032644437,0.0002146775,0.0046489206,0.001599634,0.00053791,0.0021174424,0.00063205295,0.009579548,0.022934256,0.018645795,0.9386477,0.00011558428],"about_ca_topic_score_codex":0.0070211743,"about_ca_topic_score_gemma":0.010083773,"teacher_disagreement_score":0.99930197,"about_ca_system_score_codex":0.001780331,"about_ca_system_score_gemma":0.0024516212,"threshold_uncertainty_score":0.22721195},"labels":[],"label_agreement":null},{"id":"W1482602274","doi":"10.1007/11574620_78","title":"The FungalWeb Ontology: Semantic Web Challenges in Bioinformatics and Genomics","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Semantic Web; Ontology; World Wide Web; Genomics; Information retrieval; Genome; Biology; Genetics; Gene","score_opus":0.02338673493883864,"score_gpt":0.25031716855397973,"score_spread":0.2269304336151411,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1482602274","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008114845,0.05665299,0.8132246,0.070968,0.004009301,0.00011526932,0.0024713853,0.0035551316,0.040888473],"genre_scores_gemma":[0.06103488,0.0947736,0.79783714,0.009778017,0.0027473674,0.00025179566,0.007126818,0.0013087234,0.025141643],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990551,0.00019914203,0.0001051686,0.00010784479,0.00045192265,0.00008075384],"domain_scores_gemma":[0.9974348,0.0014612874,0.000105108935,0.00035388608,0.00035089615,0.00029406088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040108473,0.00079252664,0.0010519555,0.0027734425,0.0014604499,0.007957511,0.001948816,0.0021726824,0.0036849526],"category_scores_gemma":[0.004111997,0.00060912356,0.0011643348,0.005764385,0.003648376,0.021074718,0.0028664893,0.00394715,0.002027673],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031090192,0.000072607305,0.00041259188,0.0006424285,0.000027414348,0.0002312701,0.0006513224,0.001440539,0.0014904214,0.6478271,0.095686026,0.25148723],"study_design_scores_gemma":[0.000007984071,0.000009659465,0.00026682534,0.00034788804,0.000020934984,0.0005364963,0.00044680297,0.0068343813,0.00078648725,0.56603116,0.42468733,0.000024042305],"about_ca_topic_score_codex":0.0055796765,"about_ca_topic_score_gemma":0.007699971,"teacher_disagreement_score":0.007957511,"about_ca_system_score_codex":0.0021380058,"about_ca_system_score_gemma":0.0043550357,"threshold_uncertainty_score":0.021211624},"labels":[],"label_agreement":null},{"id":"W148402895","doi":"10.3233/978-1-60750-044-5-643","title":"Computerizing Clinical Pathways: Ontology-Based Modeling and Execution","year":2009,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Ontology; Workflow; Abstraction; Software engineering; Point (geometry); Information retrieval; Programming language; Database","score_opus":0.1061958295246904,"score_gpt":0.40351191405614556,"score_spread":0.2973160845314552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W148402895","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004298376,0.0000781113,0.99160284,0.00046621574,0.000015486426,0.00022398126,0.00058955513,0.0011426831,0.001582793],"genre_scores_gemma":[0.07297376,0.00041921219,0.9232315,0.00011564751,0.000015607933,0.00041526018,0.0017217138,0.00018285193,0.0009244467],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976699,0.0007395166,0.0003602344,0.000349871,0.0007522218,0.00012821553],"domain_scores_gemma":[0.9968675,0.0018111365,0.0003254288,0.00050616026,0.00037866287,0.0001112282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003603781,0.0008815356,0.0005847413,0.0017906536,0.0010011862,0.0039498187,0.002434097,0.0009770581,0.0016317359],"category_scores_gemma":[0.0071892086,0.0007111524,0.0022959383,0.001940676,0.0012685403,0.0040777526,0.0017896045,0.0019259442,0.00048671005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011084591,0.00020119149,0.005463085,0.0005556903,0.00024734376,0.0006186334,0.0018085886,0.6200277,0.0052307374,0.23875289,0.0052773664,0.12170588],"study_design_scores_gemma":[0.000032992164,0.000027064334,0.00049032446,0.00011868898,0.00010541486,0.00014506438,0.00022199556,0.88964427,0.004798572,0.07909149,0.025283407,0.00004068122],"about_ca_topic_score_codex":0.03309423,"about_ca_topic_score_gemma":0.038622826,"teacher_disagreement_score":0.03309423,"about_ca_system_score_codex":0.0026910782,"about_ca_system_score_gemma":0.0062618563,"threshold_uncertainty_score":0.06580317},"labels":[],"label_agreement":null},{"id":"W1484050045","doi":"10.1007/11424918_34","title":"A Supervised Learning Approach to Acronym Identification","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":94,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Acronym; Computer science; Identification (biology); Task (project management); Supervised learning; Artificial intelligence; Space (punctuation); Machine learning; Artificial neural network","score_opus":0.019655489861354552,"score_gpt":0.2625764853685928,"score_spread":0.24292099550723825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1484050045","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005308127,0.00046292014,0.98775196,0.00027389315,0.0001695697,0.00015620404,0.00066796015,0.0034177778,0.0017915339],"genre_scores_gemma":[0.05487598,0.00027424048,0.93588585,0.00022198356,0.00022819766,0.00026907263,0.0025683343,0.00027540413,0.005400913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995726,0.0013667728,0.00039486302,0.0011334771,0.0012068667,0.00017205338],"domain_scores_gemma":[0.9926086,0.0037344652,0.00036296094,0.0013806267,0.0017320275,0.00018144769],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002826645,0.0011458091,0.0017699333,0.0043155323,0.0021700417,0.0023428237,0.0047909906,0.0018875197,0.005164893],"category_scores_gemma":[0.008179484,0.00076602114,0.0022476825,0.0047426983,0.0011620307,0.0037676394,0.002968189,0.0025304793,0.0033844765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001923947,0.00056589703,0.0015366877,0.00025742358,0.0001707262,0.00016490466,0.0002180941,0.028598914,0.0051282872,0.018040506,0.031972162,0.913154],"study_design_scores_gemma":[0.000055655193,0.000087741886,0.00069828826,0.00005550985,0.00008719559,0.0003201248,0.00017186302,0.8979626,0.0051924107,0.081861675,0.013453082,0.000053872838],"about_ca_topic_score_codex":0.005748968,"about_ca_topic_score_gemma":0.012913997,"teacher_disagreement_score":0.005748968,"about_ca_system_score_codex":0.0010393882,"about_ca_system_score_gemma":0.0028659906,"threshold_uncertainty_score":0.017278254},"labels":[],"label_agreement":null},{"id":"W1485005296","doi":"10.1002/bult.2015.1720410208","title":"Mapping the linguistic context of citations","year":2015,"lang":"en","type":"article","venue":"Bulletin of the Association for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Citation; Section (typography); Computer science; Context (archaeology); Linguistics; Variety (cybernetics); Natural language processing; Information retrieval; Artificial intelligence; World Wide Web; History","score_opus":0.023380419388280147,"score_gpt":0.2630822862316919,"score_spread":0.23970186684341174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1485005296","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7787164,0.019936856,0.067629695,0.020989833,0.0063024284,0.00052815874,0.025032274,0.0011092828,0.079755075],"genre_scores_gemma":[0.93226284,0.00406137,0.04414484,0.00092072465,0.002291737,0.00039778833,0.009740082,0.00033571525,0.0058449274],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99515224,0.0018846121,0.0008649177,0.0008015848,0.0011375017,0.00015905245],"domain_scores_gemma":[0.94639003,0.038471002,0.006128608,0.0013859222,0.0071900217,0.00043442068],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0038518948,0.0002745798,0.0004424972,0.013058267,0.0015840869,0.0058498723,0.0005618637,0.0005652566,0.007790489],"category_scores_gemma":[0.060253404,0.0002243278,0.00025606892,0.015351935,0.00088090077,0.0035086179,0.0021689858,0.000851295,0.0014296944],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00095489266,0.00016824945,0.1224865,0.0064559774,0.00030723758,0.002268866,0.043349914,0.0018046425,0.018223247,0.11963937,0.1216856,0.56265545],"study_design_scores_gemma":[0.00010009514,0.00019870639,0.24587567,0.0027420497,0.0003133089,0.0017769795,0.025584798,0.011401577,0.008110785,0.067329936,0.6363734,0.0001926652],"about_ca_topic_score_codex":0.0020894036,"about_ca_topic_score_gemma":0.0026166814,"teacher_disagreement_score":0.98694175,"about_ca_system_score_codex":0.0014770187,"about_ca_system_score_gemma":0.0016424162,"threshold_uncertainty_score":0.026061714},"labels":[],"label_agreement":null},{"id":"W1491247056","doi":"10.1007/3-540-45486-1_28","title":"Towards an Automated Citation Classifier","year":2000,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":98,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Classifier (UML); Computer science; Citation; Artificial intelligence; Classification scheme; Pattern recognition (psychology); Data mining; Machine learning; World Wide Web","score_opus":0.023392352280063335,"score_gpt":0.29233770733267594,"score_spread":0.26894535505261263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1491247056","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026722357,0.0032866325,0.9341973,0.0028058677,0.0012558085,0.00039309775,0.0044879923,0.01909815,0.007752751],"genre_scores_gemma":[0.101557866,0.0013965981,0.8677128,0.0007714134,0.0014033323,0.00045859118,0.009400361,0.0005338285,0.016765267],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969472,0.0005373363,0.00030988376,0.00071201555,0.0012842922,0.00020931402],"domain_scores_gemma":[0.9909524,0.0037403102,0.00035393232,0.0008743309,0.0037903416,0.00028872656],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0034672455,0.0010163201,0.002958194,0.010959267,0.0023861092,0.0051415404,0.0029762916,0.0036905115,0.0053455727],"category_scores_gemma":[0.014686856,0.00065967813,0.0017996142,0.00677029,0.0006949731,0.0055820635,0.0024956223,0.0025519598,0.008366441],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017476574,0.00044449195,0.004579974,0.00028571175,0.00014555051,0.0001693621,0.00008360104,0.008004805,0.0121158045,0.015989048,0.05045343,0.90755355],"study_design_scores_gemma":[0.0000795458,0.0001722092,0.0028616437,0.00020044114,0.00034791825,0.00061778974,0.00020093052,0.83154553,0.026366796,0.080797456,0.056733776,0.00007590205],"about_ca_topic_score_codex":0.0030942236,"about_ca_topic_score_gemma":0.0055457153,"teacher_disagreement_score":0.98904073,"about_ca_system_score_codex":0.0009807962,"about_ca_system_score_gemma":0.0028893268,"threshold_uncertainty_score":0.018336773},"labels":[],"label_agreement":null},{"id":"W1493883950","doi":"10.1007/978-3-540-78135-6_28","title":"Mixing Statistical and Symbolic Approaches for Chemical Names Recognition","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"Computer science; Matching (statistics); Classifier (UML); Pattern recognition (psychology); Artificial intelligence; Bayes' theorem; Data mining; Machine learning; Algorithm; Mathematics; Bayesian probability; Statistics","score_opus":0.043268031851047294,"score_gpt":0.26146878430102055,"score_spread":0.21820075244997325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1493883950","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011578921,0.0006297494,0.98410046,0.00032967835,0.00007265815,0.000043049185,0.00022453228,0.0015625835,0.0014584117],"genre_scores_gemma":[0.2725716,0.0009974771,0.71983206,0.00020946916,0.00027282903,0.00021932363,0.0015965122,0.00037848332,0.003922234],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99778455,0.0006962633,0.00022324524,0.0004189035,0.00072884443,0.00014809195],"domain_scores_gemma":[0.9921794,0.005675225,0.0003280535,0.00094916904,0.0007439017,0.00012429486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021465814,0.0009841505,0.0012208807,0.00343022,0.0007820286,0.002493213,0.002147881,0.0014806088,0.0041862694],"category_scores_gemma":[0.01077572,0.0006296185,0.0013202374,0.004041145,0.0013437689,0.0047484008,0.0026543606,0.0018047523,0.0018765338],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023117935,0.00016694282,0.0016892334,0.00018823394,0.0001251962,0.00009326141,0.00018865248,0.08645566,0.009265808,0.032227688,0.0029248414,0.86644334],"study_design_scores_gemma":[0.000012037652,0.000036823163,0.0005201699,0.00001690334,0.00003214044,0.00006325096,0.00007096856,0.89817125,0.0042135087,0.09501461,0.0018231858,0.000025150403],"about_ca_topic_score_codex":0.0038187765,"about_ca_topic_score_gemma":0.0080278395,"teacher_disagreement_score":0.0041862694,"about_ca_system_score_codex":0.00092206814,"about_ca_system_score_gemma":0.0017372047,"threshold_uncertainty_score":0.014004469},"labels":[],"label_agreement":null},{"id":"W1498975562","doi":"10.3389/fninf.2015.00013","title":"Text mining for neuroanatomy using WhiteText with an updated corpus and a new web application","year":2015,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Toronto","funders":"National Institute of Mental Health; National Institute of General Medical Sciences; Natural Sciences and Engineering Research Council of Canada; National Institutes of Health","keywords":"Computer science; Standardization; Information retrieval; Software; Statement (logic); Natural language processing; Neuroinformatics; Artificial intelligence; World Wide Web; Data science; Programming language","score_opus":0.027131043785827372,"score_gpt":0.26763330806017016,"score_spread":0.2405022642743428,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1498975562","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057677783,0.003233492,0.34022966,0.0058118207,0.0011227654,0.0024437932,0.43933573,0.12352461,0.02662031],"genre_scores_gemma":[0.041766163,0.0011798417,0.4185177,0.0006188809,0.00028371345,0.0036992044,0.51430863,0.007954896,0.011670965],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99787736,0.00040116467,0.00051083,0.00053259847,0.0006284407,0.000049618426],"domain_scores_gemma":[0.988395,0.00662549,0.000684966,0.0016603809,0.0022035611,0.00043067045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003054751,0.0014080334,0.00073097757,0.010662381,0.0013898636,0.002583138,0.0015167559,0.0010555027,0.02119202],"category_scores_gemma":[0.014933226,0.0006741909,0.00095833285,0.008041228,0.00088526955,0.0053343177,0.0030454902,0.0017661428,0.010345948],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004509084,0.00030696002,0.0047061164,0.004942694,0.0001852192,0.0018628839,0.0022304894,0.003278238,0.024103167,0.0110508595,0.5192601,0.42762232],"study_design_scores_gemma":[0.0002736636,0.00011769642,0.017381757,0.0006966115,0.00018291191,0.0023543036,0.00094950345,0.027952902,0.02283338,0.016309604,0.910761,0.00018662753],"about_ca_topic_score_codex":0.0042815027,"about_ca_topic_score_gemma":0.008335349,"teacher_disagreement_score":0.02119202,"about_ca_system_score_codex":0.0009757334,"about_ca_system_score_gemma":0.002037277,"threshold_uncertainty_score":0.07089436},"labels":[],"label_agreement":null},{"id":"W1499347819","doi":"10.1007/978-3-540-69828-9_15","title":"Bio2RDF : A Semantic Web Atlas of Post Genomic Knowledge about Human and Mouse","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"UniProt; Computer science; Information retrieval; Graph; World Wide Web; Atlas (anatomy); Semantic Web; Theoretical computer science; Biology","score_opus":0.014688225399841433,"score_gpt":0.2570209195131966,"score_spread":0.24233269411335517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1499347819","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017954318,0.005546575,0.44146177,0.0013760838,0.00049881364,0.00022792598,0.3566283,0.109118275,0.0671879],"genre_scores_gemma":[0.06661067,0.0070984755,0.3300077,0.0011332134,0.00023712093,0.00051017926,0.54298913,0.010615134,0.04079836],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997441,0.000029411567,0.000030246263,0.00006420819,0.00010644908,0.000025550582],"domain_scores_gemma":[0.99946934,0.00015291403,0.000071348986,0.00015731071,0.00008948296,0.000059628575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00066213205,0.00093111466,0.0005762895,0.0067383097,0.000641875,0.0017577334,0.0012210435,0.001387374,0.01796452],"category_scores_gemma":[0.001466837,0.00041093727,0.0007633021,0.0048440397,0.00040413885,0.0027710828,0.0013784136,0.0009552867,0.009087115],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006442082,0.0001947151,0.005734608,0.0031830687,0.00023018436,0.0019900082,0.0012554463,0.0072231037,0.054218657,0.12536933,0.4089649,0.39099184],"study_design_scores_gemma":[0.000029009041,0.000022051056,0.0045125624,0.0002464875,0.00009229964,0.0015638137,0.00010395069,0.0048559885,0.0109605035,0.041675787,0.9358915,0.000046043584],"about_ca_topic_score_codex":0.0048905895,"about_ca_topic_score_gemma":0.008478117,"teacher_disagreement_score":0.01796452,"about_ca_system_score_codex":0.0007445169,"about_ca_system_score_gemma":0.0014466103,"threshold_uncertainty_score":0.060097277},"labels":[],"label_agreement":null},{"id":"W1501940922","doi":"10.1007/978-3-642-03262-2_7","title":"Modeling the Form and Function of Clinical Practice Guidelines: An Ontological Model to Computerize Clinical Practice Guidelines","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Executable; Computer science; Ontology; CpG site; Information retrieval; Formalism (music); Programming language; Function (biology); Human–computer interaction; Software engineering; Theoretical computer science; DNA methylation; Gene; Chemistry","score_opus":0.1747238835813525,"score_gpt":0.4480311735038706,"score_spread":0.2733072899225181,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1501940922","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011302699,0.00054106495,0.9721067,0.004030636,0.00010315032,0.0004923348,0.0043135546,0.0011829779,0.00592694],"genre_scores_gemma":[0.08480897,0.00070025376,0.9073049,0.00044460205,0.00003766436,0.0004110136,0.004872601,0.00012944118,0.0012906172],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99733657,0.0010065287,0.0006002393,0.00032254707,0.0006013159,0.00013270396],"domain_scores_gemma":[0.99273175,0.00469211,0.00062379276,0.00069704157,0.0010018492,0.00025348744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0043914733,0.0008158922,0.00061016274,0.003680327,0.001125728,0.00476157,0.002164393,0.0017541527,0.0022788413],"category_scores_gemma":[0.016384624,0.0008356917,0.0025997793,0.0041996907,0.0013224593,0.005435446,0.0022880374,0.0021467358,0.0007437104],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016564474,0.0003267796,0.00959782,0.0010572625,0.0003357148,0.00090194395,0.004163286,0.09710222,0.0031407739,0.64635676,0.017682597,0.21916926],"study_design_scores_gemma":[0.00008116767,0.00005433828,0.0016972246,0.00068582216,0.0004806808,0.00072101405,0.0008032451,0.4091649,0.0030492498,0.4888229,0.09435539,0.00008409715],"about_ca_topic_score_codex":0.024378026,"about_ca_topic_score_gemma":0.037192617,"teacher_disagreement_score":0.024378026,"about_ca_system_score_codex":0.0025523785,"about_ca_system_score_gemma":0.0050339755,"threshold_uncertainty_score":0.048472285},"labels":[],"label_agreement":null},{"id":"W1504124766","doi":"10.1186/1471-2105-14-s3-s14","title":"Protein Function Prediction using Text-based Features extracted from the Biomedical Literature: The CAFA Challenge","year":2013,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Protein function prediction; Classifier (UML); Function (biology); DNA microarray; Gene ontology; Artificial intelligence; Data mining; Precision and recall; Protein function; Computational biology; Machine learning; Gene; Biology; Genetics","score_opus":0.019870383805722213,"score_gpt":0.2384098960389598,"score_spread":0.2185395122332376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1504124766","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49922782,0.017672509,0.16437976,0.014472689,0.0017100923,0.0020901863,0.16855825,0.122143544,0.009745201],"genre_scores_gemma":[0.39422622,0.003310341,0.41921648,0.0016432336,0.0007866656,0.001062259,0.17362821,0.0011508404,0.0049758903],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99663574,0.00050879794,0.0005389025,0.0011319093,0.0010169472,0.0001676286],"domain_scores_gemma":[0.9792244,0.013567581,0.0016710985,0.0014658584,0.003182848,0.00088822324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004754442,0.0020868082,0.0017628352,0.008255143,0.0013130577,0.002401283,0.0023713862,0.003556804,0.004086746],"category_scores_gemma":[0.020632023,0.000399575,0.0015064685,0.0044734674,0.00059879746,0.0052828896,0.0020595284,0.0018456524,0.0037159754],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014956556,0.0013506749,0.027994106,0.007044558,0.00044404727,0.0031614841,0.0011143538,0.008975372,0.07211426,0.0016123877,0.1865379,0.68815523],"study_design_scores_gemma":[0.0006993729,0.0017857845,0.11133221,0.0013967643,0.00077308615,0.011950273,0.0031755378,0.5381447,0.13508604,0.011185955,0.18387677,0.00059355184],"about_ca_topic_score_codex":0.006204196,"about_ca_topic_score_gemma":0.0068660146,"teacher_disagreement_score":0.008255143,"about_ca_system_score_codex":0.0016503761,"about_ca_system_score_gemma":0.0024287929,"threshold_uncertainty_score":0.02514416},"labels":[],"label_agreement":null},{"id":"W1504767436","doi":"10.1007/978-3-642-02976-9_10","title":"Towards the Merging of Multiple Clinical Protocols and Guidelines via Ontology-Driven Modeling","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Ontology; CpG site; Clinical Practice; Information retrieval; Software engineering; Data mining; Medicine","score_opus":0.08441899203843836,"score_gpt":0.3695302620340862,"score_spread":0.28511126999564784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1504767436","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002885707,0.00022651124,0.99103254,0.0005947567,0.000036085803,0.00038194278,0.000829835,0.002697921,0.0013147639],"genre_scores_gemma":[0.02746127,0.00028395088,0.9681484,0.00022014126,0.000016303631,0.0002777159,0.0026801284,0.00031977377,0.0005923635],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9906376,0.0032627496,0.001535802,0.001230347,0.0029842725,0.00034925717],"domain_scores_gemma":[0.9850721,0.0082663605,0.0012240078,0.002460431,0.0025411032,0.0004361332],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015567986,0.0012721041,0.0017774885,0.0048675835,0.0013937405,0.007433395,0.0048327046,0.0026046727,0.002724204],"category_scores_gemma":[0.028639017,0.0019598708,0.0056594172,0.0054551056,0.0014687099,0.00663606,0.006601494,0.0037162001,0.0014217204],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000464593,0.0005837052,0.0053341114,0.0019781871,0.0012814753,0.0018691315,0.0048206984,0.23165299,0.009768702,0.2193633,0.019848088,0.503035],"study_design_scores_gemma":[0.00010066474,0.00008490268,0.0007189894,0.000600974,0.0006211138,0.00041923847,0.0005237079,0.76189214,0.010750385,0.17026289,0.05390835,0.000116617],"about_ca_topic_score_codex":0.020076618,"about_ca_topic_score_gemma":0.030730885,"teacher_disagreement_score":0.020076618,"about_ca_system_score_codex":0.0027351321,"about_ca_system_score_gemma":0.008792003,"threshold_uncertainty_score":0.08233237},"labels":[],"label_agreement":null},{"id":"W1510421640","doi":"10.1007/978-0-387-34347-1_11","title":"Ontoligent Interactive Query Tool","year":2006,"lang":"en","type":"book-chapter","venue":"Semantic web and beyond","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Query language; Web query classification; RDF query language; Information retrieval; Query optimization; Sargable; Web search query; Query expansion; Variety (cybernetics); Syntax; Ontology; Semantic reasoner; Download; World Wide Web; Search engine; Natural language processing; Artificial intelligence","score_opus":0.008920276348439282,"score_gpt":0.23620889445504387,"score_spread":0.22728861810660458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1510421640","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006493891,0.00073205074,0.48325956,0.0008348375,0.00033986702,0.0003913,0.052675202,0.3822194,0.073053904],"genre_scores_gemma":[0.08986219,0.0016259087,0.46486259,0.00262421,0.00026644615,0.0012573313,0.21359724,0.07670008,0.14920403],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990402,0.00014031425,0.0001044394,0.00018316023,0.00045909666,0.00007279637],"domain_scores_gemma":[0.99895525,0.00053229096,0.00004301191,0.00022948954,0.00017507347,0.00006490995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011216176,0.00129666,0.00096750155,0.002384887,0.0008121464,0.0033627716,0.0013847055,0.000935188,0.07456325],"category_scores_gemma":[0.0031103033,0.00075745676,0.0008599009,0.002057924,0.0004732467,0.003931084,0.0032977187,0.0010655341,0.031808447],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069310056,0.000121872785,0.0010970847,0.0010385637,0.00009500899,0.0005144321,0.00054126047,0.0012300465,0.018435908,0.031099092,0.6733961,0.27173737],"study_design_scores_gemma":[0.000099183395,0.00003686388,0.00090111146,0.00012611046,0.000065648375,0.0007980231,0.00019128805,0.015661113,0.02680968,0.019948931,0.9352647,0.000097315955],"about_ca_topic_score_codex":0.0029208434,"about_ca_topic_score_gemma":0.003629711,"teacher_disagreement_score":0.07456325,"about_ca_system_score_codex":0.0007608366,"about_ca_system_score_gemma":0.0010589497,"threshold_uncertainty_score":0.24943894},"labels":[],"label_agreement":null},{"id":"W1510720609","doi":"10.1007/978-0-387-48438-9_14","title":"Ontology Design for Biomedical Text Mining","year":2007,"lang":"en","type":"book-chapter","venue":"Semantic Web","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Ontology; Computer science; Information retrieval; Data science; World Wide Web; Epistemology; Philosophy","score_opus":0.05595275854141806,"score_gpt":0.30267914806369517,"score_spread":0.2467263895222771,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1510720609","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014462137,0.0011344085,0.9911446,0.000718532,0.000116055904,0.00015746601,0.0008688336,0.0016777259,0.0027362094],"genre_scores_gemma":[0.025253125,0.0022137351,0.96233404,0.00045792208,0.00010098338,0.00049407635,0.004087897,0.00040797898,0.0046501867],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9979857,0.0005219887,0.00044570677,0.0002916317,0.00065100234,0.00010400472],"domain_scores_gemma":[0.9985343,0.000639591,0.000110736764,0.00027575195,0.00037233633,0.00006731222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031213344,0.00076709024,0.0010917931,0.002545139,0.0010816286,0.0032381527,0.0022750273,0.001125713,0.0044041444],"category_scores_gemma":[0.004843874,0.00074705237,0.002227301,0.0033864328,0.0012092182,0.005929844,0.0024251093,0.0019774071,0.002493247],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000115326846,0.00015575175,0.00065174425,0.0017596468,0.00020173968,0.00043135553,0.0010758619,0.009534997,0.009256059,0.4367267,0.034529887,0.50556093],"study_design_scores_gemma":[0.000040101862,0.000036728434,0.000371718,0.0004571839,0.00019475812,0.0006422472,0.00043578426,0.079174265,0.012362038,0.67902863,0.22720236,0.000054178912],"about_ca_topic_score_codex":0.002543199,"about_ca_topic_score_gemma":0.0038461904,"teacher_disagreement_score":0.0044041444,"about_ca_system_score_codex":0.0014083614,"about_ca_system_score_gemma":0.002063953,"threshold_uncertainty_score":0.016507387},"labels":[],"label_agreement":null},{"id":"W1514701801","doi":"10.1007/978-3-540-68123-6_47","title":"A Dynamic Window Based Passage Extraction Algorithm for Genomics Information Retrieval","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Genomics; Paragraph; Algorithm; Window (computing); Data mining; Focus (optics); Artificial intelligence; Information retrieval; Genome; World Wide Web; Biology","score_opus":0.011321051877579093,"score_gpt":0.2531834802643988,"score_spread":0.2418624283868197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1514701801","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018500164,0.0014376149,0.9618883,0.0001454885,0.0002133104,0.00036047233,0.0023129717,0.014029266,0.0011123434],"genre_scores_gemma":[0.042923696,0.00049964373,0.94520104,0.00007531372,0.00010512127,0.00031000987,0.006296906,0.000643351,0.003944861],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896634,0.00013812986,0.00016760739,0.00028420036,0.00035611575,0.00008770478],"domain_scores_gemma":[0.9981987,0.0008763237,0.0000869446,0.00023736407,0.00050972996,0.00009084852],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014088209,0.000946731,0.0020844075,0.005821731,0.0014799506,0.0021151167,0.0022842905,0.0011394738,0.0072757914],"category_scores_gemma":[0.0040082065,0.0005348399,0.0012433786,0.0074532633,0.00045153117,0.0030148963,0.0016479125,0.0010712919,0.004203908],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053369434,0.00023124632,0.0011524073,0.00028221137,0.00012177239,0.00019139756,0.00022401131,0.0041949586,0.030933158,0.004139926,0.02003204,0.9379632],"study_design_scores_gemma":[0.00048835645,0.0008464503,0.006264297,0.00013648227,0.0007015523,0.0017176874,0.00045615216,0.77798206,0.101560675,0.02294648,0.086676754,0.00022295307],"about_ca_topic_score_codex":0.0064937184,"about_ca_topic_score_gemma":0.006617672,"teacher_disagreement_score":0.0072757914,"about_ca_system_score_codex":0.0005882989,"about_ca_system_score_gemma":0.0014959837,"threshold_uncertainty_score":0.024339914},"labels":[],"label_agreement":null},{"id":"W1523733132","doi":"10.1089/bio.2010.0036","title":"Biospecimen Reporting for Improved Study Quality","year":2011,"lang":"en","type":"article","venue":"Biopreservation and Biobanking","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Cancer Institute; Lawrence Berkeley National Laboratory; National Institutes of Health; U.S. Department of Health and Human Services","keywords":"Quality (philosophy); Chemistry; Computer science; Physics","score_opus":0.22557409679479226,"score_gpt":0.37868458560412427,"score_spread":0.153110488809332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1523733132","genre_codex":"methods","genre_gemma":"methods","domain_codex":"reporting","domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":"reporting","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0072883614,0.0032115667,0.5681596,0.10350336,0.058727533,0.12525958,0.060719628,0.017802218,0.055328205],"genre_scores_gemma":[0.026205195,0.0036113777,0.7394064,0.028251735,0.012696746,0.1271564,0.038090304,0.0054073734,0.019174423],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.299286,0.2843877,0.32769245,0.012714414,0.06920083,0.006718666],"domain_scores_gemma":[0.05952472,0.2672043,0.15002643,0.18857585,0.32923454,0.005434134],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.5414884,0.0044569843,0.010794873,0.031507563,0.005736988,0.02238233,0.010518341,0.015096509,0.043372538],"category_scores_gemma":[0.7991002,0.005831078,0.0064521604,0.03846375,0.008697005,0.01768727,0.01353572,0.021819193,0.036267992],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018680748,0.00055743585,0.0055636484,0.017537236,0.00061742787,0.0012621923,0.00387388,0.0025112296,0.006315356,0.0293336,0.7069742,0.22358559],"study_design_scores_gemma":[0.0015914022,0.00049780007,0.013360155,0.016702384,0.0005381329,0.0011499744,0.0016349078,0.011101978,0.012290035,0.040780954,0.89965063,0.00070172676],"about_ca_topic_score_codex":0.0029126878,"about_ca_topic_score_gemma":0.002122846,"teacher_disagreement_score":0.4585116,"about_ca_system_score_codex":0.00987327,"about_ca_system_score_gemma":0.054929063,"threshold_uncertainty_score":0.56542647},"labels":[],"label_agreement":null},{"id":"W1525672751","doi":"10.1186/1471-2105-6-75","title":"Ranking the whole MEDLINE database according to a large training set using text indexing","year":2005,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Genomics","funders":"U.S. National Library of Medicine; Universität Stuttgart; Ontario Genomics Institute; Stem Cell Network; Ontario Genomics; Ontario Innovation Trust","keywords":"MEDLINE; Computer science; Information retrieval; Relevance (law); Ranking (information retrieval); Set (abstract data type); Controlled vocabulary; Vocabulary; Subject (documents); National library; Natural language processing; Artificial intelligence; World Wide Web; Library science; Linguistics","score_opus":0.07540264134123546,"score_gpt":0.33009792757829903,"score_spread":0.25469528623706356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1525672751","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6029789,0.01529621,0.2663768,0.0032113479,0.0007644142,0.0034106793,0.027630622,0.064177886,0.016153194],"genre_scores_gemma":[0.4762374,0.0033073677,0.433378,0.00089942076,0.00042826158,0.0013468164,0.07880161,0.0010976559,0.004503501],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99491453,0.0012763981,0.00083021686,0.0012926705,0.0014765309,0.00020968019],"domain_scores_gemma":[0.9768553,0.015145066,0.0007023145,0.0016398005,0.005106782,0.0005507214],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065587447,0.0022607248,0.0022724213,0.012004533,0.0021024223,0.0027839455,0.002712545,0.0022587196,0.004679297],"category_scores_gemma":[0.03719905,0.0006405343,0.0014677542,0.0074469377,0.0007812581,0.003825879,0.0019834114,0.001191104,0.0039206813],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012984201,0.0012891775,0.021950414,0.0032653955,0.00045450267,0.0016973673,0.0008803419,0.019718748,0.057365034,0.0012504266,0.041609563,0.8492205],"study_design_scores_gemma":[0.0007037664,0.0035949056,0.07097629,0.0011468179,0.0017749479,0.0049333484,0.0021364088,0.6957645,0.14027867,0.007595102,0.070745125,0.00035007807],"about_ca_topic_score_codex":0.008930406,"about_ca_topic_score_gemma":0.01041051,"teacher_disagreement_score":0.012004533,"about_ca_system_score_codex":0.0016675576,"about_ca_system_score_gemma":0.0024856238,"threshold_uncertainty_score":0.034686387},"labels":[],"label_agreement":null},{"id":"W152648670","doi":"","title":"A Methodology for Encoding Problem Lists with SNOMED CT in General Practice.","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Terminology; Computer science; Matching (statistics); General practice; Encoding (memory); Artificial intelligence; Natural language processing; Information retrieval; Medicine; Linguistics","score_opus":0.08766094539423021,"score_gpt":0.35273863970864106,"score_spread":0.26507769431441086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W152648670","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017905325,0.00015776405,0.9879462,0.00052591204,0.00009763877,0.0010568748,0.002150571,0.003901279,0.0023731946],"genre_scores_gemma":[0.0061348868,0.000083727944,0.9901136,0.00011296456,0.000015083027,0.00053238915,0.0019968401,0.00024208451,0.0007684601],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98551613,0.0061595133,0.0031131562,0.0015526382,0.0034157943,0.00024271046],"domain_scores_gemma":[0.9702857,0.013975257,0.0033253268,0.006018275,0.0059277196,0.00046777798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013998301,0.001151625,0.0006308589,0.010686218,0.0019949698,0.005030084,0.0025648982,0.0015457592,0.007239217],"category_scores_gemma":[0.04853514,0.0012152183,0.0021868516,0.009840644,0.0018973218,0.0063473424,0.005374638,0.0025981585,0.00358277],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002624355,0.00038438986,0.004253556,0.0025889925,0.00026644525,0.0013233018,0.008236751,0.011223575,0.011404545,0.1878517,0.046110895,0.72609335],"study_design_scores_gemma":[0.00014626546,0.00035029752,0.0031373512,0.0017613218,0.0002816226,0.0040118955,0.0043005925,0.08462607,0.028145416,0.23541635,0.6374516,0.00037128374],"about_ca_topic_score_codex":0.006308969,"about_ca_topic_score_gemma":0.012667544,"teacher_disagreement_score":0.013998301,"about_ca_system_score_codex":0.002081748,"about_ca_system_score_gemma":0.0068254126,"threshold_uncertainty_score":0.074030995},"labels":[],"label_agreement":null},{"id":"W1526655180","doi":"10.3233/978-1-60750-535-8-387","title":"Realism for scientific ontologies","year":2010,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Realism; Scientific realism; Epistemology; Ontology; Computer science; Philosophy","score_opus":0.05172723554712649,"score_gpt":0.3066043419196457,"score_spread":0.2548771063725192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1526655180","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063888463,0.015080374,0.41519767,0.06563376,0.002648904,0.00022417026,0.00059046806,0.00039136724,0.49384442],"genre_scores_gemma":[0.584602,0.0134743145,0.30669475,0.01768748,0.007041833,0.0014334582,0.0015207771,0.00059394934,0.06695138],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9924816,0.0035475395,0.0005597441,0.0009673023,0.0021040938,0.00033979412],"domain_scores_gemma":[0.9875179,0.0083114,0.00065143855,0.0017966522,0.0012459969,0.00047658719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010602894,0.00090681936,0.0010624339,0.002946513,0.0050662896,0.0078696,0.0019234903,0.0051264544,0.009388256],"category_scores_gemma":[0.02014605,0.00069759716,0.0014393706,0.00183913,0.019303577,0.016204217,0.00539115,0.008612853,0.0027405366],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000013570493,0.0000013267642,0.000010393131,0.000012536294,0.0000010504581,0.000009020723,0.000095017625,0.00008011487,0.0000172687,0.9977108,0.0009082214,0.0011528573],"study_design_scores_gemma":[0.0000025810907,0.0000016786253,0.000013605758,0.000026191718,0.00000131577,0.000018280338,0.00003784856,0.0002954281,0.000025532338,0.9772149,0.0223597,0.000003023924],"about_ca_topic_score_codex":0.0020990297,"about_ca_topic_score_gemma":0.0015101942,"teacher_disagreement_score":0.010602894,"about_ca_system_score_codex":0.0055348524,"about_ca_system_score_gemma":0.0024266173,"threshold_uncertainty_score":0.056074142},"labels":[],"label_agreement":null},{"id":"W1530863506","doi":"10.1161/str.46.suppl_1.tp191","title":"Abstract T P191: Using Clinical Trial Data to Generate Causative Classification System (CCS) Ischemic Stroke Phenotypes for the NINDS Stroke Genetics Network (SiGN)","year":2015,"lang":"en","type":"article","venue":"Stroke","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Subtyping; Medicine; Clinical trial; CRFS; Randomized controlled trial; Stroke (engine); Leverage (statistics); Data mining; Artificial intelligence; Pathology; Computer science","score_opus":0.2965367265954936,"score_gpt":0.4141645675446374,"score_spread":0.11762784094914375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1530863506","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28988108,0.0014449804,0.5964364,0.008834736,0.0016907909,0.02790535,0.04388052,0.015237759,0.014688341],"genre_scores_gemma":[0.4356442,0.00025837994,0.53810275,0.00089667144,0.00030951572,0.007339331,0.015317522,0.00093805656,0.0011936161],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9601299,0.024336604,0.007141607,0.0033918477,0.004660687,0.0003394035],"domain_scores_gemma":[0.64941096,0.2761195,0.020948244,0.026342649,0.02535042,0.0018282064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06343349,0.00089938473,0.00091090665,0.00463558,0.0008387124,0.0036507277,0.0016627674,0.0009000984,0.0063706026],"category_scores_gemma":[0.30085823,0.0005070275,0.0017937312,0.0027745531,0.0006117807,0.001142793,0.0015885888,0.0017884423,0.0016739634],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0056129275,0.00053639617,0.21797569,0.0039299782,0.0016292298,0.001756758,0.0028569042,0.03491011,0.004658427,0.012383107,0.15028335,0.56346714],"study_design_scores_gemma":[0.007100772,0.0030685707,0.17479113,0.0038412246,0.0024215963,0.0023259488,0.0012322569,0.59495705,0.03301884,0.06948386,0.10729184,0.0004669499],"about_ca_topic_score_codex":0.002830572,"about_ca_topic_score_gemma":0.0035296082,"teacher_disagreement_score":0.06343349,"about_ca_system_score_codex":0.0012590848,"about_ca_system_score_gemma":0.0066749323,"threshold_uncertainty_score":0.33547235},"labels":[],"label_agreement":null},{"id":"W1533751446","doi":"","title":"Ontology based holonic diagnostic system (OHDS) for the research and control of unknown diseases","year":2005,"lang":"en","type":"article","venue":"eSpace (Curtin University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Ontology; Computer science; Scalability; Annotation; Upper ontology; Software engineering; Data science; Information retrieval; Semantic Web; Artificial intelligence; Database","score_opus":0.022568147100265338,"score_gpt":0.27293042362523906,"score_spread":0.2503622765249737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1533751446","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006344835,0.0005941046,0.96448255,0.0018069831,0.00030222582,0.00036403435,0.0012639739,0.015387661,0.009453713],"genre_scores_gemma":[0.19639589,0.000989137,0.7889312,0.001444222,0.00020675466,0.00048253094,0.00481786,0.00043794027,0.0062945182],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99808407,0.0004841403,0.0003024629,0.00036514344,0.0006577849,0.000106531385],"domain_scores_gemma":[0.9953811,0.0017951243,0.000432575,0.00102749,0.0010179363,0.00034571267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004295703,0.00076313847,0.0005504146,0.0037667914,0.0010187654,0.0025578775,0.0015301292,0.0010295676,0.0040199687],"category_scores_gemma":[0.006563298,0.00023976272,0.0009738615,0.0013411574,0.0010358341,0.003345043,0.0028863803,0.0011824099,0.0014747818],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005992812,0.00025659052,0.011228034,0.0010922917,0.0002912522,0.0013849903,0.0013928391,0.011105338,0.024054367,0.16551025,0.053265404,0.7298193],"study_design_scores_gemma":[0.00023601869,0.0002735962,0.0070652696,0.00067445537,0.0006833776,0.0021703069,0.00087338797,0.28152955,0.054451603,0.23625706,0.41550273,0.00028270783],"about_ca_topic_score_codex":0.0037227625,"about_ca_topic_score_gemma":0.00335577,"teacher_disagreement_score":0.004295703,"about_ca_system_score_codex":0.0015161898,"about_ca_system_score_gemma":0.0031112735,"threshold_uncertainty_score":0.022718132},"labels":[],"label_agreement":null},{"id":"W1534640868","doi":"10.1007/978-3-642-02879-3_11","title":"Slicing through the Scientific Literature","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Information retrieval; Syntax; Knowledge base; Domain (mathematical analysis); Knowledge representation and reasoning; Scientific literature; Representation (politics); Natural language; Slicing; World Wide Web; Natural language processing; Artificial intelligence","score_opus":0.015981989477162412,"score_gpt":0.26467784056824145,"score_spread":0.24869585109107903,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1534640868","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023540962,0.014081856,0.90922004,0.0047740135,0.0011001283,0.00053447427,0.008833409,0.00844337,0.029471777],"genre_scores_gemma":[0.058829315,0.007223794,0.9093894,0.00048487366,0.00031809072,0.00016565608,0.01306247,0.001353349,0.009173069],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968329,0.0005718073,0.00048560553,0.0006329436,0.0013548735,0.000121893616],"domain_scores_gemma":[0.9828773,0.007940953,0.00092961185,0.0030844377,0.004634139,0.0005336474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047105066,0.0010320948,0.0011881884,0.016079584,0.0015232407,0.0053969747,0.0017525877,0.0006690479,0.009852874],"category_scores_gemma":[0.01598116,0.000870644,0.002482174,0.011100416,0.0018487269,0.008080981,0.0044530826,0.0015388466,0.004089666],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013451684,0.000058062804,0.0028848655,0.0019963544,0.00024460934,0.0006930349,0.0027045899,0.0030925896,0.008391139,0.13397339,0.048270565,0.79755634],"study_design_scores_gemma":[0.000043413078,0.000095575975,0.002687773,0.0014644128,0.0005680487,0.001429449,0.003039792,0.02920479,0.02094086,0.438359,0.5020766,0.000090373425],"about_ca_topic_score_codex":0.004729828,"about_ca_topic_score_gemma":0.008956305,"teacher_disagreement_score":0.016079584,"about_ca_system_score_codex":0.0012474385,"about_ca_system_score_gemma":0.0056058536,"threshold_uncertainty_score":0.03296119},"labels":[],"label_agreement":null},{"id":"W1538444001","doi":"10.3414/me13-02-0029","title":"Semi Automated Transformation to OWL Formatted Files as an Approach to Data Integration","year":2014,"lang":"en","type":"article","venue":"Methods of Information in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Population and Public Health","funders":"National Institutes of Health; National Institute for Health and Care Research; NIHR Biomedical Research Centre, Royal Marsden NHS Foundation Trust/Institute of Cancer Research","keywords":"Computer science; Metadata; Information retrieval; Web Ontology Language; Ontology; Interoperability; Correctness; Linked data; Metadata repository; World Wide Web; Semantic Web; Database; Programming language","score_opus":0.05014347448942495,"score_gpt":0.41415133912738117,"score_spread":0.3640078646379562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1538444001","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006833427,0.000070096474,0.92668843,0.0004374513,0.00014745697,0.00096420496,0.006755245,0.054760907,0.0033427035],"genre_scores_gemma":[0.04718343,0.00012652719,0.9221261,0.0003099034,0.00004361276,0.001006361,0.017202048,0.008966685,0.003035251],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9924965,0.0018616096,0.0013334021,0.0012750147,0.0027096414,0.00032385823],"domain_scores_gemma":[0.9760878,0.013235648,0.0011770653,0.0053502945,0.0038334336,0.00031570395],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008440058,0.0016854649,0.00087110087,0.0047948323,0.0012863023,0.005025571,0.002917342,0.0015275286,0.010980421],"category_scores_gemma":[0.027138585,0.001445097,0.0036994745,0.0038521783,0.0017743018,0.0045891954,0.005756673,0.0033049253,0.0048359297],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014565531,0.0012725897,0.0077834665,0.0028718277,0.00062521367,0.004661021,0.011782467,0.052905124,0.045723822,0.08779006,0.08400961,0.69911826],"study_design_scores_gemma":[0.00033273696,0.0003707495,0.003482318,0.000982521,0.00027846007,0.0019080181,0.0030280089,0.30564338,0.13401115,0.12232387,0.42721164,0.00042723867],"about_ca_topic_score_codex":0.006762128,"about_ca_topic_score_gemma":0.005401874,"teacher_disagreement_score":0.010980421,"about_ca_system_score_codex":0.002000221,"about_ca_system_score_gemma":0.0028901622,"threshold_uncertainty_score":0.044635832},"labels":[],"label_agreement":null},{"id":"W1542243371","doi":"","title":"Automatic acquisition of long-distance acronym definitions","year":2003,"lang":"en","type":"article","venue":"Hybrid Intelligent Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Acronym; Computer science; Treebank; Lexicon; Component (thermodynamics); Natural language processing; Artificial intelligence; Matching (statistics); Bigram; Identification (biology); Support vector machine; Word (group theory); Linguistics; Trigram; Parsing","score_opus":0.0323483849835466,"score_gpt":0.27025896935624577,"score_spread":0.23791058437269919,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1542243371","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26580057,0.003695911,0.66888964,0.0014696202,0.00054240413,0.00046949356,0.0108148735,0.026775451,0.021542048],"genre_scores_gemma":[0.39854723,0.0011700324,0.5722519,0.00020605575,0.0001912517,0.00023308204,0.018888282,0.0010456268,0.0074665323],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881977,0.0003038555,0.00014295317,0.00038071672,0.00029185068,0.00006082658],"domain_scores_gemma":[0.99571544,0.0022783934,0.00063505943,0.0003643318,0.00091385,0.000092901806],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008397307,0.0006418102,0.00080454885,0.0046824403,0.0007273635,0.0017242308,0.00082722655,0.00066395226,0.0053886985],"category_scores_gemma":[0.007258241,0.00044609397,0.0005599679,0.0028779334,0.0003670216,0.004585851,0.0016309834,0.00089190865,0.0036412072],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002736417,0.00015587277,0.012057764,0.0010576926,0.000113560134,0.0008728965,0.0012918715,0.0016952885,0.08476199,0.008445961,0.03577643,0.8534969],"study_design_scores_gemma":[0.00014677629,0.00039256658,0.08809213,0.000546296,0.00037227265,0.007872963,0.0053198338,0.31665358,0.2180598,0.038779795,0.32343292,0.00033110174],"about_ca_topic_score_codex":0.0009999579,"about_ca_topic_score_gemma":0.0022338617,"teacher_disagreement_score":0.0053886985,"about_ca_system_score_codex":0.00036151282,"about_ca_system_score_gemma":0.0006762987,"threshold_uncertainty_score":0.018027008},"labels":[],"label_agreement":null},{"id":"W1543782282","doi":"10.1089/omi.2006.10.185","title":"National Center for Biomedical Ontology: Advancing Biomedicine through Structured Organization of Scientific Knowledge","year":2006,"lang":"en","type":"article","venue":"OMICS A Journal of Integrative Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"National Human Genome Research Institute","keywords":"Ontology; Biomedicine; Computer science; Dissemination; Open Biomedical Ontologies; Context (archaeology); Resource (disambiguation); Knowledge management; Data science; Upper ontology; World Wide Web; Quality (philosophy); Semantic Web; Suggested Upper Merged Ontology; Bioinformatics","score_opus":0.013108609816200769,"score_gpt":0.3132670196152181,"score_spread":0.3001584097990173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1543782282","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020762351,0.0043689017,0.92219365,0.031068664,0.0017356586,0.0010616361,0.003114033,0.007128542,0.027252756],"genre_scores_gemma":[0.011773712,0.004391577,0.9669339,0.004220407,0.00053947436,0.001127804,0.0076276585,0.00071730197,0.0026681111],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.980053,0.008909091,0.0032418019,0.0020106405,0.005060806,0.0007246522],"domain_scores_gemma":[0.94936335,0.018199835,0.003813125,0.017606169,0.0070987004,0.00391878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03531606,0.0013944681,0.0018337832,0.011176434,0.004745195,0.013098968,0.005682231,0.0042138128,0.0059603285],"category_scores_gemma":[0.057579488,0.0014279548,0.0028749562,0.013505346,0.0069392533,0.020009201,0.020504255,0.007828604,0.003958399],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008408076,0.00018160683,0.0018270985,0.0016777435,0.00018151644,0.00026179204,0.0028404063,0.0018480683,0.0026615518,0.6664125,0.10555577,0.21646796],"study_design_scores_gemma":[0.0000870769,0.0000472579,0.0010444929,0.0013649852,0.000114969574,0.00027183612,0.00072902715,0.007883747,0.0016465978,0.47122428,0.5155025,0.000083323604],"about_ca_topic_score_codex":0.013944199,"about_ca_topic_score_gemma":0.010852256,"teacher_disagreement_score":0.03531606,"about_ca_system_score_codex":0.0061773793,"about_ca_system_score_gemma":0.036794256,"threshold_uncertainty_score":0.18677145},"labels":[],"label_agreement":null},{"id":"W1554930632","doi":"10.1111/j.1540-8159.2010.02867.x","title":"Unusual Marker Annotation: Something Shifty is Going On","year":2010,"lang":"en","type":"article","venue":"Pacing and Clinical Electrophysiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Medicine; University hospital; Emergency department; General hospital; Library science; Family medicine; Psychiatry","score_opus":0.014119574148154037,"score_gpt":0.3389112193559577,"score_spread":0.3247916452078037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1554930632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016731642,0.00655185,0.7120023,0.22183657,0.009029429,0.0002695692,0.0039294777,0.013775742,0.015873468],"genre_scores_gemma":[0.17312478,0.0057538445,0.7195388,0.056184866,0.006319627,0.00040346335,0.010947777,0.007156192,0.020570619],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9904512,0.0039575957,0.0011512088,0.0017459448,0.0021301548,0.0005639586],"domain_scores_gemma":[0.93080217,0.030521434,0.0026425542,0.017871914,0.014886726,0.0032752405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019289395,0.0013599738,0.0022975272,0.005520216,0.00440126,0.011237846,0.0047990624,0.004352262,0.01808093],"category_scores_gemma":[0.058266025,0.0010081183,0.0020839197,0.0048143105,0.00639758,0.045503985,0.008698966,0.010933506,0.005389826],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000551143,0.00028234607,0.007387192,0.0017309373,0.00022464267,0.0008180682,0.0065137977,0.0014368762,0.009567425,0.13145591,0.21354324,0.6264884],"study_design_scores_gemma":[0.000061516716,0.00007634567,0.003321092,0.0018644888,0.00022982174,0.0015130348,0.008583072,0.011881065,0.0048540346,0.39718497,0.5701344,0.00029614376],"about_ca_topic_score_codex":0.0052275546,"about_ca_topic_score_gemma":0.0068079676,"teacher_disagreement_score":0.019289395,"about_ca_system_score_codex":0.0017427199,"about_ca_system_score_gemma":0.005615527,"threshold_uncertainty_score":0.10201323},"labels":[],"label_agreement":null},{"id":"W1555519040","doi":"10.3233/sw-2011-0048","title":"Taking flight with OWL2","year":2011,"lang":"en","type":"article","venue":"Semantic Web","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Aeronautics; Computer science; Engineering","score_opus":0.026822638090422856,"score_gpt":0.24049680674359167,"score_spread":0.21367416865316882,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1555519040","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060682646,0.0014351245,0.73228467,0.020697085,0.0056221914,0.00053740805,0.0117928265,0.07322626,0.14833623],"genre_scores_gemma":[0.07328981,0.002129465,0.7843494,0.01032015,0.0010004706,0.00037698957,0.030360281,0.024209915,0.07396351],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958676,0.00080369826,0.00031521908,0.0005549811,0.002146771,0.00031174696],"domain_scores_gemma":[0.9952154,0.001149713,0.00018495657,0.0015848149,0.0015011051,0.00036390757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060313037,0.0010036812,0.00061101845,0.0017534022,0.0023631589,0.006531321,0.002189019,0.0015582094,0.032064762],"category_scores_gemma":[0.015389218,0.0011338746,0.0026665728,0.0012924746,0.0013788897,0.0114220595,0.005520554,0.0058611333,0.016562076],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002803508,0.00019663229,0.0015718421,0.00047500862,0.00022801338,0.0007885153,0.0013936836,0.002465576,0.0043781013,0.19368768,0.5081112,0.28642336],"study_design_scores_gemma":[0.000057264915,0.000033936947,0.00027012642,0.00016164032,0.000048598115,0.00029715517,0.00040500629,0.006632621,0.0025022833,0.11455632,0.87497395,0.00006107194],"about_ca_topic_score_codex":0.018441843,"about_ca_topic_score_gemma":0.018883927,"teacher_disagreement_score":0.032064762,"about_ca_system_score_codex":0.0012922568,"about_ca_system_score_gemma":0.0030278538,"threshold_uncertainty_score":0.10726732},"labels":[],"label_agreement":null},{"id":"W1556136921","doi":"10.1109/iembs.2003.1280923","title":"CellMapBase-an information system supporting high-throughput proteomics for the Cell Map project","year":2004,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Set (abstract data type); Process (computing); Software engineering; Data science; World Wide Web; Programming language","score_opus":0.015410915083074258,"score_gpt":0.2696143202313335,"score_spread":0.2542034051482593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1556136921","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013800854,0.0016186485,0.52469486,0.0018791341,0.00073773705,0.0015793125,0.1036454,0.33519867,0.01684551],"genre_scores_gemma":[0.05771241,0.0024435124,0.53554136,0.0017781212,0.00037140373,0.0024579247,0.36709997,0.022898965,0.009696318],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985997,0.00019964267,0.00017336459,0.00023896493,0.00068188546,0.00010645335],"domain_scores_gemma":[0.9946451,0.0016297582,0.0005270903,0.0012241483,0.0010008083,0.00097309373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0056178253,0.0014388635,0.0017467737,0.005945819,0.0020777222,0.006706763,0.0043845423,0.0023845446,0.010976083],"category_scores_gemma":[0.011533215,0.0011667713,0.0008017123,0.0070994864,0.00068645563,0.006622999,0.0039750063,0.0026778486,0.0137802055],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015235342,0.0007615985,0.0066276956,0.0020035696,0.0003664373,0.0014028766,0.001253986,0.0053395513,0.03892887,0.039796088,0.6554978,0.24649794],"study_design_scores_gemma":[0.00065164594,0.00023031632,0.0074938564,0.00041939178,0.0002569286,0.0015379064,0.00043691482,0.06608747,0.056817282,0.048475564,0.8171481,0.00044458907],"about_ca_topic_score_codex":0.0037676163,"about_ca_topic_score_gemma":0.0031793418,"teacher_disagreement_score":0.010976083,"about_ca_system_score_codex":0.0011604804,"about_ca_system_score_gemma":0.0040140813,"threshold_uncertainty_score":0.036718607},"labels":[],"label_agreement":null},{"id":"W1563809394","doi":"","title":"An author by any other name","year":2000,"lang":"en","type":"article","venue":"PubMed Central","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Data science; World Wide Web; Library science; Information retrieval","score_opus":0.01171903226891437,"score_gpt":0.2487321806372628,"score_spread":0.23701314836834841,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1563809394","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009146862,0.0060880124,0.0038566787,0.017057994,0.06862442,0.00010256995,0.0017497494,0.0019359161,0.8996698],"genre_scores_gemma":[0.0054300516,0.0020070423,0.0018254855,0.0040943855,0.0037922729,0.00005689611,0.0010192863,0.000803006,0.9809715],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99813217,0.0003742026,0.00017051151,0.0004496231,0.00068451767,0.00018901698],"domain_scores_gemma":[0.98997647,0.0015981386,0.0007451231,0.0022014878,0.0028978428,0.0025808914],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002006373,0.0009922514,0.001005629,0.0017898388,0.0018144308,0.0061563984,0.0014538275,0.0021424734,0.6988106],"category_scores_gemma":[0.011647329,0.00033628193,0.00063499936,0.002646542,0.0009479829,0.0044259927,0.0040659783,0.002631719,0.7467484],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000029180323,0.000033179924,0.00034034747,0.00026443694,0.000012044773,0.0002400101,0.00013722693,0.00003141005,0.00062786345,0.006944085,0.8685208,0.12281955],"study_design_scores_gemma":[0.0000022685786,0.000006011006,0.000084237174,0.000084820895,0.0000024613144,0.00020051627,0.00003794583,0.000018122486,0.00009627916,0.0006708267,0.9987924,0.000004110205],"about_ca_topic_score_codex":0.000431584,"about_ca_topic_score_gemma":0.0010358449,"teacher_disagreement_score":0.99799365,"about_ca_system_score_codex":0.0007972278,"about_ca_system_score_gemma":0.0017468676,"threshold_uncertainty_score":0.42961013},"labels":[],"label_agreement":null},{"id":"W1564137013","doi":"10.1007/978-3-642-15120-0_10","title":"Algorithm for Grounding Mutation Mentions from Text to Protein Sequences","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Mutation; Context (archaeology); Identifier; Task (project management); Annotation; Precision and recall; Natural language processing; Computational biology; Information retrieval; Artificial intelligence; Data mining; Genetics; Biology; Programming language; Gene","score_opus":0.018840528726906752,"score_gpt":0.2836457604943696,"score_spread":0.26480523176746285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1564137013","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037987035,0.00054633257,0.91767496,0.00096568844,0.00038545384,0.0010133676,0.006955683,0.03164828,0.0028231486],"genre_scores_gemma":[0.060595192,0.0001392999,0.9242714,0.00025866585,0.000057730736,0.00032818125,0.009944337,0.0005013228,0.003903887],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998988,0.00009947667,0.00011363528,0.00040449225,0.00030575963,0.00008868783],"domain_scores_gemma":[0.9966695,0.0018927914,0.00019649563,0.00038678947,0.0007529563,0.000101495236],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015180606,0.0016045009,0.0014025578,0.005003571,0.0013064202,0.002177071,0.0028530376,0.00263438,0.010885854],"category_scores_gemma":[0.005564632,0.0008900363,0.0022054017,0.00315165,0.00085381255,0.0027587938,0.0022230367,0.001650289,0.004544198],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004750435,0.00035614075,0.0052412343,0.00039329246,0.00024230451,0.0003819896,0.00024883766,0.020528965,0.010825191,0.006287958,0.028643986,0.9263752],"study_design_scores_gemma":[0.00038148556,0.000342647,0.0029560674,0.00022672594,0.00054007897,0.0008120406,0.0005068784,0.87637365,0.032512896,0.05268551,0.032594174,0.00006784514],"about_ca_topic_score_codex":0.006574345,"about_ca_topic_score_gemma":0.014607926,"teacher_disagreement_score":0.010885854,"about_ca_system_score_codex":0.0010690327,"about_ca_system_score_gemma":0.0028561403,"threshold_uncertainty_score":0.03641683},"labels":[],"label_agreement":null},{"id":"W156651503","doi":"10.3233/978-1-60750-581-5-213","title":"Ontologies, Semantic Technologies, and Intelligence: Looking Toward the Future","year":2010,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Computer science; Semantic technology; Data science; Semantic Web; World Wide Web; Semantic computing","score_opus":0.023347511693418952,"score_gpt":0.27228667099408754,"score_spread":0.2489391593006686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W156651503","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020768081,0.30267048,0.039273266,0.06346696,0.0065609007,0.0000402039,0.00015299526,0.0003283505,0.58542997],"genre_scores_gemma":[0.08359331,0.39504194,0.07902713,0.02952805,0.008370025,0.00015167652,0.0008548305,0.0005764821,0.40285665],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9991677,0.00026803036,0.000041381318,0.00009438885,0.00035731163,0.00007128943],"domain_scores_gemma":[0.99928147,0.00045685575,0.00002518707,0.000052377844,0.00012788549,0.00005623762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016219744,0.00091529277,0.00062722183,0.0022311234,0.001422993,0.0076298234,0.0008187867,0.001827401,0.0068404237],"category_scores_gemma":[0.001481015,0.00036411316,0.0005524518,0.0036356393,0.0070453114,0.020561423,0.0016952141,0.004781626,0.0034381594],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000006018224,0.000010674483,0.00004254372,0.00015317308,0.0000057524194,0.000034610366,0.00080549147,0.00025660865,0.00014396431,0.9094112,0.048320618,0.040809408],"study_design_scores_gemma":[0.0000018479662,0.0000032677624,0.000072952666,0.00021859631,0.0000032843523,0.000082149396,0.0006036472,0.00031767756,0.00007989057,0.38829637,0.6103137,0.0000065627846],"about_ca_topic_score_codex":0.0035182475,"about_ca_topic_score_gemma":0.0050995713,"teacher_disagreement_score":0.0076298234,"about_ca_system_score_codex":0.0039301957,"about_ca_system_score_gemma":0.0028072551,"threshold_uncertainty_score":0.028515697},"labels":[],"label_agreement":null},{"id":"W1574574070","doi":"10.1007/978-3-642-12275-0_8","title":"Improving Medical Information Retrieval with PICO Element Detection","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Information retrieval; Task (project management); Process (computing); Element (criminal law); Population; Data mining","score_opus":0.006390804664680461,"score_gpt":0.227661113417714,"score_spread":0.22127030875303355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1574574070","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09165611,0.006484125,0.88083285,0.00062430045,0.0005521999,0.00036693487,0.0026354387,0.00990063,0.006947394],"genre_scores_gemma":[0.192165,0.0014898431,0.79223484,0.0003715492,0.0004100487,0.00023440634,0.005177119,0.00038000653,0.007537252],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99907255,0.00013934732,0.00009055094,0.0001653161,0.0004623961,0.00006991231],"domain_scores_gemma":[0.9981394,0.0010091637,0.000089079615,0.00019893388,0.00050171243,0.000061560226],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007665502,0.0007182097,0.0013178606,0.0052687535,0.0005132271,0.0014531479,0.0010253253,0.00086044346,0.0039370954],"category_scores_gemma":[0.0039618677,0.00037210167,0.00080478424,0.0034096956,0.0003027531,0.0018797365,0.0014387238,0.000541663,0.0025175847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047927353,0.00016622394,0.0037023528,0.0003309117,0.00010207637,0.00012964399,0.000091917675,0.0023280191,0.043889035,0.0019618878,0.011091571,0.93572706],"study_design_scores_gemma":[0.00018526119,0.0006481259,0.011992573,0.00017867917,0.00062242785,0.0028880977,0.0003988495,0.7234717,0.19121641,0.015187168,0.05307995,0.00013069954],"about_ca_topic_score_codex":0.002235651,"about_ca_topic_score_gemma":0.0031369317,"teacher_disagreement_score":0.0052687535,"about_ca_system_score_codex":0.0003332143,"about_ca_system_score_gemma":0.0006667617,"threshold_uncertainty_score":0.013170898},"labels":[],"label_agreement":null},{"id":"W1575535538","doi":"10.18438/b8gk7m","title":"PubMed is Slightly More Sensitive but More Time Intensive to Search than Ovid MEDLINE","year":2011,"lang":"en","type":"article","venue":"Evidence Based Library and Information Practice","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"Mount Royal University","funders":"","keywords":"MEDLINE; Medicine; Randomized controlled trial; Cochrane Library; National library; Online search; Controlled vocabulary; Information retrieval; Computer science; Library science; Surgery","score_opus":0.02584876114246174,"score_gpt":0.27202300988235373,"score_spread":0.246174248739892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1575535538","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007986377,0.3268698,0.03298643,0.08492725,0.023560334,0.052855022,0.32573083,0.009060581,0.13602346],"genre_scores_gemma":[0.039320473,0.29485497,0.25754672,0.08710749,0.01360821,0.13519308,0.13072848,0.005558655,0.036081973],"study_design_codex":"systematic_review","study_design_gemma":"observational","domain_scores_codex":[0.8600018,0.058232233,0.057122067,0.0045895935,0.018602226,0.0014520959],"domain_scores_gemma":[0.484432,0.38235244,0.053269543,0.014361675,0.061290905,0.004293414],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07333309,0.0024594446,0.014339206,0.07852007,0.001536558,0.012264279,0.0050293133,0.004186379,0.16266754],"category_scores_gemma":[0.3651858,0.0016037357,0.004668547,0.0766124,0.002156961,0.011848916,0.0068681417,0.0029263687,0.041209776],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012127938,0.000104395425,0.0014525357,0.4700986,0.002173986,0.0005493766,0.0014447253,0.00028538203,0.0016017117,0.004168248,0.3098598,0.20704843],"study_design_scores_gemma":[0.001134927,0.00036399663,0.008286324,0.23422058,0.003289777,0.00087414746,0.0013067458,0.0003187395,0.00068336,0.0060511893,0.7432331,0.00023710757],"about_ca_topic_score_codex":0.0042409324,"about_ca_topic_score_gemma":0.008182688,"teacher_disagreement_score":0.9266669,"about_ca_system_score_codex":0.0046618246,"about_ca_system_score_gemma":0.020771245,"threshold_uncertainty_score":0.5441772},"labels":[],"label_agreement":null},{"id":"W1579787587","doi":"","title":"Dataset: OA Publication rates","year":2013,"lang":"br","type":"dataset","venue":"uO Research (University of Ottawa)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Information retrieval; Computer science","score_opus":0.0785297309123752,"score_gpt":0.35164851949261144,"score_spread":0.27311878858023625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1579787587","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00019089582,0.0000926975,0.00003998594,0.000067801346,0.000017195032,0.000009787515,0.9990068,0.0001562607,0.0004186041],"genre_scores_gemma":[0.000473488,0.00009432369,0.00020094578,0.00003683964,0.000011778846,0.00004515431,0.9985531,0.000038959057,0.00054541556],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965089,0.00043909153,0.0008521685,0.0007397084,0.0010579366,0.00040205132],"domain_scores_gemma":[0.98993886,0.002589076,0.0019092916,0.0016042806,0.0030019935,0.0009566431],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.0018580217,0.0017002327,0.0019072156,0.009482067,0.0008885209,0.0035067499,0.0024957887,0.0019746884,0.04484139],"category_scores_gemma":[0.014943456,0.0006105336,0.0014218674,0.02035829,0.0004656927,0.001695239,0.0017483842,0.0018036711,0.06465427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012112357,0.00003257802,0.0023057293,0.0011420147,0.00006352012,0.000032538654,0.00003239888,0.00025910785,0.00018900902,0.00061117596,0.99097764,0.0042331773],"study_design_scores_gemma":[0.0002479535,0.000034629164,0.014747666,0.0005008784,0.00007848997,0.0001332929,0.00012125998,0.00048895954,0.00037873638,0.00094808027,0.9822688,0.000051223804],"about_ca_topic_score_codex":0.030306878,"about_ca_topic_score_gemma":0.05899945,"teacher_disagreement_score":0.998142,"about_ca_system_score_codex":0.002851437,"about_ca_system_score_gemma":0.004975046,"threshold_uncertainty_score":0.15000933},"labels":[],"label_agreement":null},{"id":"W1580895464","doi":"10.3233/ao-2011-0082","title":"The RNA Ontology (RNAO): An ontology for integrating RNA sequence and structure data","year":2011,"lang":"en","type":"article","venue":"Applied Ontology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Ontology; Sequence (biology); Information retrieval; RNA; Biology; Gene","score_opus":0.08717297104569614,"score_gpt":0.32226851980320614,"score_spread":0.23509554875751,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1580895464","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061843945,0.0022956282,0.95213336,0.0025723327,0.000393071,0.00040132052,0.014492125,0.00716468,0.014363222],"genre_scores_gemma":[0.049677633,0.0057928655,0.9056522,0.0018242347,0.0003126379,0.000875428,0.028353384,0.001203099,0.0063085877],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99770015,0.00045984884,0.0004024575,0.0003721353,0.0009406753,0.00012472476],"domain_scores_gemma":[0.9972963,0.0012761457,0.00043508742,0.0004127422,0.00042417372,0.00015546652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028102386,0.0009330293,0.0010032922,0.005697543,0.001559916,0.0029154574,0.0018187474,0.0012697931,0.0020787853],"category_scores_gemma":[0.0059045777,0.0004775606,0.0018774441,0.0055798646,0.0014773062,0.006639062,0.0023940133,0.0019020187,0.0014074503],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019101276,0.00019256305,0.0048974496,0.00258743,0.00035538958,0.0012976709,0.0024876948,0.007961314,0.032346305,0.52252525,0.06407359,0.36108434],"study_design_scores_gemma":[0.00005507067,0.000061986095,0.0026661188,0.0009843514,0.00025721328,0.0012689269,0.00060830504,0.030469446,0.012210567,0.19605674,0.75520915,0.00015215423],"about_ca_topic_score_codex":0.007193628,"about_ca_topic_score_gemma":0.0073441477,"teacher_disagreement_score":0.007193628,"about_ca_system_score_codex":0.0018382592,"about_ca_system_score_gemma":0.004933629,"threshold_uncertainty_score":0.01486212},"labels":[],"label_agreement":null},{"id":"W1582205572","doi":"10.18438/b8zp4v","title":"New Search Strategies Successfully Optimize Retrieval of Clinically Sound Treatment Studies in EMBASE","year":2007,"lang":"en","type":"article","venue":"Evidence Based Library and Information Practice","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"MEDLINE; Test (biology); Medicine; Subject (documents); Set (abstract data type); Medical physics; Computer science; Information retrieval; Medical education; Family medicine; Library science","score_opus":0.056940214559720585,"score_gpt":0.38727574853573044,"score_spread":0.33033553397600984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1582205572","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.092644684,0.47250617,0.21137649,0.03925407,0.0044988617,0.10788054,0.03224633,0.0051962687,0.034396604],"genre_scores_gemma":[0.1264657,0.1002391,0.6833848,0.005781487,0.0020550105,0.06751556,0.0111878095,0.00082678546,0.0025438932],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.7100847,0.14964661,0.11767729,0.0052980385,0.016124006,0.001169325],"domain_scores_gemma":[0.2657659,0.6481123,0.034800753,0.015915781,0.034031436,0.0013738937],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.1952975,0.003913063,0.00888359,0.077795625,0.0018919214,0.009065552,0.004160406,0.0030600182,0.01876827],"category_scores_gemma":[0.6017552,0.002611095,0.0066835154,0.05470966,0.0018398698,0.014741248,0.0068113143,0.002257627,0.005078165],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029561494,0.00041119428,0.008505753,0.2689498,0.0074714483,0.0013254009,0.0050390833,0.0017920362,0.0051169563,0.004139741,0.0412337,0.65305877],"study_design_scores_gemma":[0.020017289,0.0049139177,0.05883888,0.4408687,0.094283365,0.005358424,0.009182806,0.022428714,0.013518533,0.0442141,0.28427944,0.0020959491],"about_ca_topic_score_codex":0.0023879474,"about_ca_topic_score_gemma":0.014347135,"teacher_disagreement_score":0.8047025,"about_ca_system_score_codex":0.0051795803,"about_ca_system_score_gemma":0.014305872,"threshold_uncertainty_score":0.99234146},"labels":[],"label_agreement":null},{"id":"W1585341798","doi":"10.1007/11553939_171","title":"Leximancer Concept Mapping of Patient Case Studies","year":2005,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":68,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Semantic mapping; Information retrieval; Semantic analysis (machine learning); Natural language; Natural language processing; Data science","score_opus":0.028233692346901887,"score_gpt":0.2895871211396911,"score_spread":0.2613534287927892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1585341798","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.097972706,0.0030800013,0.8652375,0.00159626,0.00034055475,0.00090498605,0.011951867,0.0058703893,0.013045845],"genre_scores_gemma":[0.34115717,0.0010727571,0.6302251,0.00032326317,0.000107790154,0.00064216624,0.021668231,0.0004887619,0.0043147635],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9947225,0.0020541581,0.00078281795,0.000980753,0.0011999803,0.00025975538],"domain_scores_gemma":[0.9871896,0.008878678,0.00045038803,0.0018165878,0.0014606455,0.00020412693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0064390814,0.00079606817,0.00088138523,0.008387775,0.0010516845,0.002706713,0.0023891542,0.0010949684,0.008059079],"category_scores_gemma":[0.023237986,0.0004154532,0.0015883197,0.0056847157,0.0007919427,0.0029110154,0.0035163707,0.0013228272,0.0015325338],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006211378,0.00032901694,0.015989492,0.0011429398,0.00039142402,0.0014498397,0.0010556064,0.023714893,0.0037270575,0.027066799,0.024101637,0.90041023],"study_design_scores_gemma":[0.00023470749,0.00040093125,0.014428295,0.00082972826,0.00077003066,0.006040018,0.0027923407,0.5858267,0.03256536,0.26819155,0.08778766,0.0001327212],"about_ca_topic_score_codex":0.003497552,"about_ca_topic_score_gemma":0.006123042,"teacher_disagreement_score":0.008387775,"about_ca_system_score_codex":0.0011362719,"about_ca_system_score_gemma":0.0028352488,"threshold_uncertainty_score":0.034053504},"labels":[],"label_agreement":null},{"id":"W1586231883","doi":"10.1096/fasebj.21.5.a400-a","title":"Creation of a Retrospective Searchable Neuropathologic Database from Print Archives: The UHN Experience","year":2007,"lang":"en","type":"article","venue":"The FASEB Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Toronto Western Hospital; University Health Network; University of Toronto; Dalhousie University; Toronto General Hospital; York University","funders":"","keywords":"Computer science; Medical diagnosis; Categorization; Prioritization; Database; Information retrieval; MEDLINE; Medicine; World Wide Web; Library science; Pathology; Artificial intelligence","score_opus":0.022467263213671384,"score_gpt":0.2945269854574984,"score_spread":0.272059722243827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1586231883","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.918864,0.0031681217,0.025201697,0.0016659211,0.00011984197,0.004359565,0.033581667,0.0009887195,0.012050562],"genre_scores_gemma":[0.8619756,0.003863968,0.07429252,0.0007410825,0.00017789286,0.0020530259,0.053147834,0.00043450104,0.003313594],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99532616,0.0014332113,0.0011423075,0.00070670375,0.0011558251,0.0002358543],"domain_scores_gemma":[0.98131126,0.007062405,0.0016249279,0.005007859,0.0037508793,0.001242637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012591498,0.00026797983,0.00055199803,0.0069906423,0.00086799834,0.0019107296,0.001380151,0.00047166666,0.003468826],"category_scores_gemma":[0.028511168,0.0004171207,0.00031171791,0.006513511,0.00057532935,0.0028084943,0.0022217287,0.000513137,0.0014052725],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017334011,0.0010811266,0.47726038,0.0013437483,0.00012445585,0.0074865436,0.008352819,0.0009916105,0.005142605,0.002042772,0.023463856,0.47097665],"study_design_scores_gemma":[0.00049772346,0.0018218196,0.71772707,0.0020450896,0.0006455632,0.024481501,0.020940132,0.012021871,0.019235613,0.0025894474,0.19768877,0.00030546752],"about_ca_topic_score_codex":0.008863606,"about_ca_topic_score_gemma":0.013939866,"teacher_disagreement_score":0.012591498,"about_ca_system_score_codex":0.0016078432,"about_ca_system_score_gemma":0.003384692,"threshold_uncertainty_score":0.066590965},"labels":[],"label_agreement":null},{"id":"W1595714231","doi":"10.1186/1472-6947-7-3","title":"A UMLS-based spell checker for natural language processing in vaccine safety","year":2007,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Public Health Agency of Canada","funders":"Oak Ridge Institute for Science and Education; Centers for Disease Control and Prevention; U.S. Department of Energy","keywords":"Spelling; Computer science; Unified Medical Language System; Spell; Natural language processing; Artificial intelligence; Linguistics","score_opus":0.017111740455615716,"score_gpt":0.348760040074018,"score_spread":0.3316482996184023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1595714231","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030351575,0.0010629662,0.7359292,0.000679206,0.0005133597,0.0016809324,0.010673423,0.21696404,0.0021452797],"genre_scores_gemma":[0.081574626,0.00032067054,0.8948398,0.0004729424,0.00013481878,0.0009167511,0.015364066,0.004791726,0.001584664],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98086596,0.006274282,0.0043525365,0.0044570076,0.003634108,0.00041611487],"domain_scores_gemma":[0.8978937,0.06902414,0.009103922,0.0073778317,0.015746992,0.00085337064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013105246,0.003116359,0.0017883121,0.007636791,0.0013844328,0.0034913803,0.0027714674,0.0015945284,0.009348561],"category_scores_gemma":[0.069957696,0.001499113,0.0029208045,0.002881855,0.0015158383,0.0044945246,0.0036324887,0.0026310184,0.006758257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016402568,0.00072811014,0.020048702,0.0041607297,0.0006725293,0.0019927658,0.002933712,0.026359608,0.055595472,0.009185081,0.07129503,0.805388],"study_design_scores_gemma":[0.0008097761,0.0008721823,0.012847042,0.0011952944,0.00072661653,0.0027238913,0.00093066186,0.607166,0.22451171,0.023467997,0.12416751,0.00058140384],"about_ca_topic_score_codex":0.009430145,"about_ca_topic_score_gemma":0.008652628,"teacher_disagreement_score":0.013105246,"about_ca_system_score_codex":0.0024164002,"about_ca_system_score_gemma":0.007308394,"threshold_uncertainty_score":0.06930798},"labels":[],"label_agreement":null},{"id":"W1598005774","doi":"10.5772/13560","title":"Social and Semantic Web Technologies for the Text-to-Knowledge Translation Process in Biomedicine","year":2011,"lang":"en","type":"book-chapter","venue":"InTech eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; Data science; Knowledge extraction; Semantic Web; Biomedicine; Knowledge integration; Process (computing); Social Semantic Web; Semantics (computer science); Domain knowledge; Knowledge base; Information retrieval; World Wide Web; Knowledge management; Artificial intelligence","score_opus":0.0657135054077038,"score_gpt":0.32333226215083355,"score_spread":0.25761875674312973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1598005774","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045131124,0.08082485,0.6661011,0.032562703,0.0035403997,0.00052764406,0.00322445,0.0053725285,0.2033333],"genre_scores_gemma":[0.052754927,0.12503974,0.7227743,0.006458364,0.002413344,0.001162598,0.0076344176,0.001690043,0.080072284],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99771094,0.0010885088,0.00021231278,0.00019035551,0.00070753857,0.000090274494],"domain_scores_gemma":[0.99656147,0.0024572636,0.00015886247,0.00043092007,0.00028551728,0.00010600552],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003824885,0.0012667653,0.0008935844,0.0060400884,0.001581991,0.007440174,0.0018581898,0.003548715,0.02093754],"category_scores_gemma":[0.004922543,0.000652796,0.0014459112,0.01146463,0.0032976372,0.017071014,0.0047216914,0.0029964491,0.013136983],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035778132,0.00007027393,0.0002635955,0.0024398032,0.000082039856,0.00066852005,0.0013355358,0.0021187805,0.0023174812,0.485907,0.08612175,0.41863948],"study_design_scores_gemma":[0.000012922595,0.000014718697,0.0003839638,0.00094699115,0.000027603137,0.00049156154,0.00073052297,0.0049564303,0.0012302266,0.50200576,0.48916844,0.000030867384],"about_ca_topic_score_codex":0.0014841569,"about_ca_topic_score_gemma":0.0023221846,"teacher_disagreement_score":0.02093754,"about_ca_system_score_codex":0.0018307815,"about_ca_system_score_gemma":0.0022090499,"threshold_uncertainty_score":0.07004303},"labels":[],"label_agreement":null},{"id":"W1598794494","doi":"10.1002/dvg.22873","title":"Xenbase: Core features, data acquisition, and data processing","year":2015,"lang":"en","type":"article","venue":"genesis","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; Biotechnology and Biological Sciences Research Council; National Institutes of Health; Wellcome Trust","keywords":"Core (optical fiber); Computer science; Data acquisition; Telecommunications; Operating system","score_opus":0.13029213013074667,"score_gpt":0.3598876030192796,"score_spread":0.2295954728885329,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1598794494","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050351606,0.0014223313,0.11089967,0.0009764694,0.00041490173,0.0022366385,0.5801455,0.28167716,0.01719214],"genre_scores_gemma":[0.014746394,0.0011118266,0.12305735,0.00094436825,0.00017452605,0.004345273,0.8265142,0.024158081,0.004947992],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964669,0.0004380383,0.00079681346,0.0007655769,0.0012729497,0.00025977672],"domain_scores_gemma":[0.9941719,0.0011567903,0.0004893634,0.0014970357,0.0020618653,0.00062302704],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051760357,0.002617798,0.002336833,0.0063102194,0.0014007398,0.0063440446,0.006019357,0.0011974444,0.04022465],"category_scores_gemma":[0.016796872,0.0015054904,0.0011245202,0.007540009,0.00078778114,0.0042262403,0.005742384,0.0022241392,0.050867304],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012418518,0.00015767459,0.004593137,0.002406652,0.00015114862,0.00051217544,0.0004000091,0.0023070427,0.0073597813,0.0045895423,0.89058214,0.08569869],"study_design_scores_gemma":[0.00030688298,0.000110158675,0.005767057,0.0005749586,0.00012981151,0.0006154274,0.00023201887,0.010304131,0.021604368,0.009865609,0.9503086,0.00018093841],"about_ca_topic_score_codex":0.005486862,"about_ca_topic_score_gemma":0.003947583,"teacher_disagreement_score":0.04022465,"about_ca_system_score_codex":0.0014431318,"about_ca_system_score_gemma":0.0036580542,"threshold_uncertainty_score":0.13456482},"labels":[],"label_agreement":null},{"id":"W1600834905","doi":"10.1186/1471-2105-6-78","title":"CGMIM: Automated text-mining of Online Mendelian Inheritance in Man (OMIM) to identify genetically-associated cancers and candidate genes","year":2005,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"Michael Smith Health Research BC","keywords":"OMIM : Online Mendelian Inheritance in Man; CDKN2A; Cancer; Gene; Mendelian inheritance; Biology; Genetics; Candidate gene; Computational biology; Bioinformatics; Phenotype","score_opus":0.020657280603231,"score_gpt":0.3130826807367318,"score_spread":0.29242540013350077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1600834905","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14531863,0.002850515,0.34396848,0.0022535808,0.00044323003,0.003698173,0.293784,0.19814014,0.009543291],"genre_scores_gemma":[0.1022668,0.0010407532,0.6842239,0.00053873414,0.000263787,0.0022161992,0.20263533,0.004027626,0.0027869313],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979668,0.0005103051,0.00039794104,0.0007026901,0.00034138543,0.00008072059],"domain_scores_gemma":[0.9878312,0.0080949385,0.0018104934,0.0009933438,0.00089129416,0.00037873967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044888128,0.001851922,0.0012394015,0.011970907,0.0008136193,0.001680872,0.0017755734,0.00094957656,0.012144503],"category_scores_gemma":[0.018004075,0.0006044346,0.0016459423,0.006036262,0.00044287406,0.001436278,0.0018481867,0.000754341,0.004495848],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028491386,0.00070495374,0.09342497,0.0055010845,0.0015346216,0.004517383,0.0023662215,0.008360124,0.033807572,0.0065141413,0.22918841,0.6112314],"study_design_scores_gemma":[0.0019594196,0.0013168658,0.19202846,0.0017670087,0.0019318786,0.011334677,0.0018974029,0.34182325,0.07702866,0.039318893,0.3291217,0.00047174725],"about_ca_topic_score_codex":0.0026298524,"about_ca_topic_score_gemma":0.0043395255,"teacher_disagreement_score":0.012144503,"about_ca_system_score_codex":0.00086336344,"about_ca_system_score_gemma":0.0020665575,"threshold_uncertainty_score":0.04062742},"labels":[],"label_agreement":null},{"id":"W160147681","doi":"10.1007/978-3-642-21043-3_30","title":"Improving Phenotype Name Recognition","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Unified Medical Language System; Natural language processing; Ontology; Named-entity recognition; Recall; Artificial intelligence; Precision and recall; Information retrieval; Biomedical text mining; Text mining; Linguistics; Task (project management)","score_opus":0.024681606249192467,"score_gpt":0.24380072161456437,"score_spread":0.2191191153653719,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W160147681","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11299655,0.0014963123,0.81525254,0.00091096235,0.0010142889,0.00013663016,0.005657131,0.050923754,0.011611786],"genre_scores_gemma":[0.32313508,0.00085891056,0.6302443,0.00068967044,0.00030897817,0.0001333643,0.019847116,0.0030407265,0.02174186],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992083,0.00012411577,0.000059793318,0.0002785918,0.00025011203,0.000079154444],"domain_scores_gemma":[0.9977411,0.0007500082,0.00012661099,0.0006839833,0.0006146431,0.00008372082],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00069029507,0.0011007105,0.0012683172,0.0017317359,0.0005351288,0.0014163838,0.0014175713,0.0010243532,0.0079951165],"category_scores_gemma":[0.002896007,0.00029177425,0.0011865976,0.0014221484,0.00027083888,0.001896149,0.0012768668,0.0011923237,0.0076525705],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029853356,0.0002540287,0.00666264,0.0001859253,0.00009346147,0.00032834904,0.000054942695,0.007298222,0.0762682,0.0026229888,0.022674154,0.8832586],"study_design_scores_gemma":[0.000100641104,0.00033863162,0.016583215,0.00008886477,0.0003949472,0.002066372,0.00023437869,0.68574756,0.21576546,0.024801854,0.053755146,0.00012297336],"about_ca_topic_score_codex":0.0016623487,"about_ca_topic_score_gemma":0.0028754072,"teacher_disagreement_score":0.0079951165,"about_ca_system_score_codex":0.00035726052,"about_ca_system_score_gemma":0.00052511215,"threshold_uncertainty_score":0.026746392},"labels":[],"label_agreement":null},{"id":"W1645387934","doi":"","title":"Las secuencias signatura de diversas proteínas demuestran la divergencia tardía del Orden Aquificales","year":2004,"lang":"es","type":"article","venue":"International Microbiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Biology; Zoology","score_opus":0.007602497835870883,"score_gpt":0.2582383090827884,"score_spread":0.25063581124691753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1645387934","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9405855,0.0032179581,0.039684836,0.0011224062,0.00049689267,0.000058668993,0.003572147,0.0025905336,0.008671034],"genre_scores_gemma":[0.94510674,0.0022465105,0.040446483,0.00023367196,0.00012400771,0.000068098736,0.003375111,0.00040195597,0.007997337],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99909604,0.000094512776,0.00011533625,0.00024672173,0.0003365565,0.00011089611],"domain_scores_gemma":[0.9959727,0.0010887989,0.00079417974,0.00057134166,0.0012606655,0.00031230898],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019489446,0.0009573983,0.0005902903,0.0025146024,0.0007682805,0.0019650888,0.00038340196,0.0006774017,0.0030314429],"category_scores_gemma":[0.0062560104,0.00029200243,0.0009267313,0.0025422703,0.0008140152,0.0017076724,0.00079870026,0.0012091357,0.0013919809],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023082597,0.0003535138,0.28226155,0.0014615031,0.0004863732,0.002066143,0.0016276112,0.003432108,0.40583923,0.011762551,0.0046438677,0.28375733],"study_design_scores_gemma":[0.000091434784,0.0005195182,0.5623033,0.00028825595,0.0010171917,0.0067578456,0.0029577794,0.02357568,0.31947342,0.02025681,0.06252147,0.00023722237],"about_ca_topic_score_codex":0.0036676931,"about_ca_topic_score_gemma":0.0040864586,"teacher_disagreement_score":0.0036676931,"about_ca_system_score_codex":0.0011052638,"about_ca_system_score_gemma":0.0011092272,"threshold_uncertainty_score":0.010307074},"labels":[],"label_agreement":null},{"id":"W165859259","doi":"10.3233/978-1-61499-488-6-286","title":"Optimizing the Efficacy of Multimedia Consumer Health Information","year":2015,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Multimedia; Computer science; Animation; Leverage (statistics); Popularity; Presentation (obstetrics); Health information; Interactive media; Multimedia information retrieval; Health care; Psychology; Medicine; Artificial intelligence","score_opus":0.05425945348917232,"score_gpt":0.3673455797746466,"score_spread":0.3130861262854743,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W165859259","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83666897,0.0016868631,0.092936985,0.0050280886,0.00010345106,0.0010049734,0.0002847074,0.00093354064,0.06135251],"genre_scores_gemma":[0.97127795,0.00031888668,0.02602902,0.00019337637,0.000048797712,0.00011711109,0.000112796835,0.00004299829,0.0018591568],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963981,0.002119686,0.00017979534,0.00028596394,0.0008402038,0.00017611764],"domain_scores_gemma":[0.96923447,0.025020834,0.0017057221,0.0012623832,0.0023001714,0.00047637892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068389303,0.00045386812,0.0004784053,0.0016190436,0.00047478345,0.0026108613,0.0007577636,0.0009309842,0.0048163235],"category_scores_gemma":[0.045806963,0.00021829283,0.00031661254,0.00084316236,0.00058287865,0.002861888,0.0013412788,0.0005970098,0.0005931841],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033389924,0.0028245358,0.03071354,0.0011775411,0.00018889438,0.00028344756,0.0036126897,0.024316054,0.019673077,0.032721996,0.0057694428,0.8753799],"study_design_scores_gemma":[0.0011654017,0.0063517606,0.09647565,0.0008898826,0.0015320969,0.0006583251,0.009419376,0.62914747,0.08279489,0.120670415,0.050603144,0.00029161558],"about_ca_topic_score_codex":0.0012340688,"about_ca_topic_score_gemma":0.00089543813,"teacher_disagreement_score":0.0068389303,"about_ca_system_score_codex":0.001215187,"about_ca_system_score_gemma":0.0010468418,"threshold_uncertainty_score":0.036168158},"labels":[],"label_agreement":null},{"id":"W1667664355","doi":"10.1111/cobi.12034","title":"Scientific Alert: The Medium is the Message","year":2013,"lang":"en","type":"article","venue":"Conservation Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Queen's University","funders":"","keywords":"Presentation (obstetrics); Variety (cybernetics); Media studies; Library science; Sociology; History; Computer science; Artificial intelligence; Medicine","score_opus":0.0265501338653203,"score_gpt":0.2744067480772579,"score_spread":0.2478566142119376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1667664355","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001983051,0.03652278,0.006000315,0.3734887,0.13437957,0.00018845605,0.0019582468,0.0037371083,0.44174176],"genre_scores_gemma":[0.03478256,0.026116448,0.004318217,0.1380978,0.058928926,0.0002721527,0.0013994559,0.0022828416,0.7338016],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99603075,0.0012057379,0.00018297456,0.00039797512,0.0018959435,0.00028662066],"domain_scores_gemma":[0.9902345,0.0030538822,0.00093810784,0.0009342473,0.0030844708,0.0017547285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032685956,0.0010400465,0.0008327045,0.0017396037,0.0025199435,0.01584145,0.0011771724,0.005228093,0.24160278],"category_scores_gemma":[0.022132125,0.0005630479,0.000637635,0.0014273749,0.0027251018,0.015567365,0.004788627,0.0068587447,0.19755168],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003449058,0.000011495012,0.000088766574,0.00021016588,0.000009248269,0.00005190573,0.00038890424,0.00001270546,0.00045527937,0.005802963,0.9404698,0.052464187],"study_design_scores_gemma":[0.00000667932,0.000013474543,0.00012387721,0.00014333572,0.0000047745416,0.00006987945,0.00029987283,0.000018370047,0.00013419906,0.0034060553,0.9957709,0.000008591036],"about_ca_topic_score_codex":0.0008876865,"about_ca_topic_score_gemma":0.0009608627,"teacher_disagreement_score":0.24160278,"about_ca_system_score_codex":0.001815325,"about_ca_system_score_gemma":0.0025568898,"threshold_uncertainty_score":0.80824184},"labels":[],"label_agreement":null},{"id":"W169573124","doi":"","title":"P224-M Management Systems and Storage of Laboratory Information at Laval University Hospital Research Center Genomic Sequencing and Genotyping Platform","year":2007,"lang":"en","type":"article","venue":"PubMed Central","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Natural Sciences and Engineering Research Council of Canada","funders":"","keywords":"Computer science; Upload; Database; Genotyping; Sample (material); Operating system; World Wide Web; Biology","score_opus":0.017287850764879823,"score_gpt":0.22989605161657695,"score_spread":0.21260820085169713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W169573124","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07462279,0.0020373801,0.24424388,0.0061772307,0.0022468043,0.00741721,0.10405764,0.39496782,0.16422926],"genre_scores_gemma":[0.2341448,0.0019462185,0.3737271,0.0051493905,0.0033954892,0.01371872,0.18766792,0.018086001,0.16216435],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934496,0.0016423377,0.00062375085,0.0018766435,0.0017781052,0.00062955223],"domain_scores_gemma":[0.9856623,0.0022404124,0.0018467677,0.0042505586,0.0036701243,0.0023299006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007000222,0.0013312072,0.0014522936,0.0049690777,0.0015257108,0.00552618,0.0044693854,0.0009298158,0.14139602],"category_scores_gemma":[0.014754967,0.00093207724,0.0004957197,0.0032843486,0.00078002026,0.0024166028,0.0036655208,0.001666972,0.11757975],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024548261,0.0003697476,0.010999733,0.00032798242,0.00008520935,0.00043439944,0.00061294646,0.0009059072,0.015423543,0.005107921,0.66328156,0.2999963],"study_design_scores_gemma":[0.0011431228,0.000948646,0.02194913,0.00029621136,0.00013056485,0.0007828099,0.00033397463,0.020548837,0.040304806,0.003838786,0.9094388,0.0002843171],"about_ca_topic_score_codex":0.0041556535,"about_ca_topic_score_gemma":0.0017313714,"teacher_disagreement_score":0.14139602,"about_ca_system_score_codex":0.0024406975,"about_ca_system_score_gemma":0.0051618144,"threshold_uncertainty_score":0.4730168},"labels":[],"label_agreement":null},{"id":"W1706534725","doi":"","title":"Recognizing named entities in biomedical texts","year":2008,"lang":"en","type":"dissertation","venue":"Summit (Simon Fraser University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Named-entity recognition; Biomedical text mining; Artificial intelligence; Natural language processing; Word (group theory); Domain (mathematical analysis); Named entity; Text mining; Task (project management); Linguistics","score_opus":0.01387467025524078,"score_gpt":0.23490134396290788,"score_spread":0.2210266737076671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1706534725","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17885065,0.0018314755,0.7869042,0.0020492529,0.00039061697,0.0005460231,0.006258481,0.012578812,0.010590488],"genre_scores_gemma":[0.27608153,0.0008994366,0.7002914,0.00040339184,0.00020105024,0.00026818743,0.015847454,0.0004348139,0.005572716],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984201,0.00046138125,0.00020180084,0.0006039525,0.00025041384,0.000062255116],"domain_scores_gemma":[0.99247426,0.0050791316,0.00061382435,0.0007544151,0.0009325867,0.000145867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020882753,0.00062703656,0.0005245176,0.0028303203,0.00067195855,0.001585442,0.0009673361,0.0013088844,0.0028804701],"category_scores_gemma":[0.0105982935,0.00038294846,0.00077308796,0.00159062,0.0005555922,0.004603778,0.0012479073,0.001175603,0.0032251887],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042391333,0.0002411261,0.016493343,0.0010408942,0.00013558261,0.0013031177,0.0016747368,0.034532756,0.08227183,0.017532814,0.024958076,0.8193919],"study_design_scores_gemma":[0.000048994334,0.00022827207,0.018031463,0.00025805537,0.0001415172,0.002132889,0.0018105807,0.7200672,0.1254251,0.04272048,0.08901522,0.000120244025],"about_ca_topic_score_codex":0.0012550134,"about_ca_topic_score_gemma":0.002885741,"teacher_disagreement_score":0.0028804701,"about_ca_system_score_codex":0.0005833491,"about_ca_system_score_gemma":0.00074579765,"threshold_uncertainty_score":0.011044025},"labels":[],"label_agreement":null},{"id":"W1757553993","doi":"10.1007/978-3-319-24027-5_29","title":"Is Concept Mapping Useful for Biomedical Information Retrieval?","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Information retrieval; Artificial intelligence","score_opus":0.03895289143602098,"score_gpt":0.28723734228522013,"score_spread":0.24828445084919915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1757553993","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03693778,0.041851457,0.8411203,0.024555625,0.002789303,0.00018616236,0.0031507672,0.00398365,0.04542491],"genre_scores_gemma":[0.40420067,0.028854826,0.5474687,0.0031133194,0.0019069619,0.0002793303,0.0035573572,0.00079096685,0.009827777],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997223,0.0013344035,0.00022219506,0.00042936736,0.00069011375,0.00010096072],"domain_scores_gemma":[0.98675364,0.010444243,0.0003688477,0.0012761351,0.0010085667,0.00014846669],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054841926,0.000780925,0.0011896709,0.005753205,0.0007550637,0.004605594,0.0017042867,0.0017410849,0.008508615],"category_scores_gemma":[0.03181604,0.0005362411,0.0011330853,0.0072382437,0.0017003769,0.015970418,0.0017806531,0.0012602499,0.0038669782],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022376639,0.00010633267,0.0022762958,0.0012029076,0.00017965565,0.00020594479,0.0005693535,0.0023226843,0.003280524,0.088213824,0.028539747,0.87287897],"study_design_scores_gemma":[0.000054217253,0.00009392528,0.0022797324,0.00049323955,0.00021344669,0.0016030823,0.0009905616,0.059831373,0.007254606,0.8391209,0.08797834,0.00008648017],"about_ca_topic_score_codex":0.001242731,"about_ca_topic_score_gemma":0.0008457107,"teacher_disagreement_score":0.008508615,"about_ca_system_score_codex":0.0006489454,"about_ca_system_score_gemma":0.0009410542,"threshold_uncertainty_score":0.02900356},"labels":[],"label_agreement":null},{"id":"W175787552","doi":"10.3233/978-1-61499-203-5-257","title":"Knowledge Translation in eHealth: Building a Virtual Community","year":2013,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Canadian Institutes of Health Research","keywords":"eHealth; Knowledge translation; Knowledge management; Computer science; Action (physics); Virtual community; Translation (biology); World Wide Web; Health care; The Internet; Political science","score_opus":0.10112606278470736,"score_gpt":0.39958247186211804,"score_spread":0.29845640907741067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W175787552","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18028578,0.0017402631,0.59656024,0.048003703,0.0012837262,0.0025119581,0.00017293,0.0013231909,0.16811825],"genre_scores_gemma":[0.7748412,0.00057703693,0.20633739,0.0019204194,0.00024955062,0.0010499833,0.00022413049,0.00020921884,0.014591118],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9688301,0.025330901,0.0007037913,0.0017445517,0.0025258188,0.0008647444],"domain_scores_gemma":[0.95699495,0.026994843,0.002137406,0.005576142,0.0024825125,0.0058140624],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031173002,0.00053805945,0.0005557243,0.0034804114,0.008577064,0.010870263,0.002710839,0.0034164495,0.008764564],"category_scores_gemma":[0.036807746,0.00048472013,0.0010793338,0.001973345,0.011466884,0.0167468,0.023882382,0.0028372582,0.0017319856],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037939698,0.001908129,0.008930007,0.0011750883,0.00018215856,0.0018816122,0.1609939,0.0055763223,0.006844778,0.30401313,0.019660456,0.488455],"study_design_scores_gemma":[0.00035083413,0.000784602,0.003958814,0.001054614,0.00013441531,0.0009889085,0.06759549,0.025770482,0.0040471307,0.5234236,0.37167287,0.00021817572],"about_ca_topic_score_codex":0.0015809746,"about_ca_topic_score_gemma":0.0015957594,"teacher_disagreement_score":0.031173002,"about_ca_system_score_codex":0.0025212083,"about_ca_system_score_gemma":0.007784878,"threshold_uncertainty_score":0.1648606},"labels":[],"label_agreement":null},{"id":"W1769350208","doi":"10.1111/jpc.12705","title":"Phenotyping: Targeting genotype's rich cousin for diagnosis","year":2014,"lang":"en","type":"review","venue":"Journal of Paediatrics and Child Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children; University of Toronto","funders":"National Health and Medical Research Council","keywords":"Medicine; Cousin; Workflow; Precision medicine; Disease; Medical genetics; Bioinformatics; Pathology; Genetics; Computer science; Gene","score_opus":0.029191982548428697,"score_gpt":0.3430395726512138,"score_spread":0.3138475901027851,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1769350208","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00025367184,0.9916952,0.0023329463,0.0027125438,0.00043214927,0.000010729946,0.000076646065,0.000043476135,0.002442582],"genre_scores_gemma":[0.0030969481,0.99100894,0.0029250851,0.001225111,0.0004871556,0.000014199192,0.00015094428,0.000010052694,0.001081613],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992824,0.00022161886,0.0001011262,0.00012387913,0.00023496765,0.00003603562],"domain_scores_gemma":[0.9978612,0.0015721398,0.00014265957,0.000068173315,0.00028407553,0.0000717397],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017483508,0.0007946455,0.0012984205,0.0043458343,0.0003366306,0.0013134417,0.0012988904,0.0015271548,0.003054173],"category_scores_gemma":[0.004151165,0.00030516033,0.0008836603,0.0025942568,0.0011135797,0.0023077389,0.0010597586,0.0023537131,0.0021645464],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004309401,0.000030697,0.0006512617,0.008274042,0.00013958478,0.00040885367,0.000121973135,0.00028066194,0.0005809553,0.010725255,0.033497747,0.94524586],"study_design_scores_gemma":[0.000011633087,0.0000309389,0.0017181312,0.0072208205,0.00017067199,0.0045778835,0.00015388007,0.0002483727,0.00053577765,0.011988365,0.97331536,0.000028103748],"about_ca_topic_score_codex":0.00232821,"about_ca_topic_score_gemma":0.0035300325,"teacher_disagreement_score":0.0043458343,"about_ca_system_score_codex":0.00095338636,"about_ca_system_score_gemma":0.0015051486,"threshold_uncertainty_score":0.010217249},"labels":[],"label_agreement":null},{"id":"W1775135849","doi":"10.1136/amiajnl-2013-002411","title":"Learning regular expressions for clinical text classification","year":2014,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":117,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; U.S. Department of Veterans Affairs","keywords":"Computer science; Natural language processing; Artificial intelligence","score_opus":0.025299992043553883,"score_gpt":0.3547212481779366,"score_spread":0.3294212561343827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1775135849","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.070952654,0.0014846586,0.902138,0.0016403936,0.00022395323,0.00087296136,0.004749591,0.015126772,0.0028111238],"genre_scores_gemma":[0.29447773,0.00066969317,0.68861866,0.00057483965,0.00023966879,0.001267897,0.011699555,0.00068463007,0.0017672542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99001896,0.0031221323,0.0016074426,0.00252906,0.0023539183,0.0003684689],"domain_scores_gemma":[0.9704927,0.019257586,0.0037602137,0.00203507,0.0040272567,0.00042724077],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007466298,0.0013325086,0.001081518,0.004541108,0.0007455472,0.002135276,0.0019356986,0.0010705058,0.0020822107],"category_scores_gemma":[0.03618214,0.0004904255,0.0011832893,0.0030309325,0.0014053476,0.0029954738,0.0015962526,0.0018366115,0.002428413],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009023743,0.00036133043,0.020750064,0.0011623712,0.00017728719,0.000973513,0.001138646,0.0262085,0.019842992,0.011828872,0.025094917,0.89155906],"study_design_scores_gemma":[0.0001625788,0.00042930365,0.0075912327,0.0004502051,0.0001764867,0.001376867,0.00076796825,0.85438204,0.04722799,0.058748517,0.028573854,0.00011301938],"about_ca_topic_score_codex":0.001858388,"about_ca_topic_score_gemma":0.0017158327,"teacher_disagreement_score":0.007466298,"about_ca_system_score_codex":0.001539932,"about_ca_system_score_gemma":0.002339389,"threshold_uncertainty_score":0.03948605},"labels":[],"label_agreement":null},{"id":"W1779982606","doi":"10.1197/jamia.m2996","title":"Towards Automatic Recognition of Scientifically Rigorous Clinical Research Evidence","year":2008,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":108,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Concordia University","funders":"National Institutes of Health","keywords":"Computer science; Artificial intelligence; Machine learning; Support vector machine; Precision and recall; Gold standard (test); Naive Bayes classifier; Classifier (UML); Information retrieval; Pattern recognition (psychology); Medicine","score_opus":0.12454115338744977,"score_gpt":0.4284704772819773,"score_spread":0.30392932389452754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1779982606","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09524005,0.024859408,0.847002,0.0059807277,0.00078829593,0.0016232871,0.0049168332,0.012795648,0.006793719],"genre_scores_gemma":[0.16564658,0.0029309953,0.82451046,0.0005464619,0.00048158754,0.00037338072,0.004658719,0.00015700565,0.0006948614],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.971629,0.010769818,0.0059169927,0.0032486527,0.0078020254,0.0006334544],"domain_scores_gemma":[0.8217102,0.105263054,0.020829711,0.011817852,0.03848493,0.0018941283],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.030663399,0.0013549016,0.002702199,0.029308757,0.0012528367,0.009236764,0.0024573992,0.002944698,0.0014683071],"category_scores_gemma":[0.12849982,0.000718268,0.0019127777,0.010323426,0.0011157724,0.005357355,0.0038116265,0.0029739975,0.002002604],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033022056,0.00028703795,0.0217776,0.0027257828,0.00035482607,0.00037194695,0.0006561795,0.004376606,0.021806492,0.00511411,0.012132443,0.93006665],"study_design_scores_gemma":[0.00035016454,0.0010590579,0.0631709,0.0037961875,0.0019094055,0.0031315575,0.0023245763,0.61117476,0.12961672,0.10229317,0.08073217,0.0004413623],"about_ca_topic_score_codex":0.002373386,"about_ca_topic_score_gemma":0.0045904135,"teacher_disagreement_score":0.9693366,"about_ca_system_score_codex":0.0016398436,"about_ca_system_score_gemma":0.008040714,"threshold_uncertainty_score":0.16216552},"labels":[],"label_agreement":null},{"id":"W178313090","doi":"10.1090/dimacs/061/14","title":"How good can a consensus get? Assessing the reliability of consensus trees in phylogenetic studies","year":2003,"lang":"en","type":"book-chapter","venue":"DIMACS series in discrete mathematics and theoretical computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Phylogenetic tree; Reliability (semiconductor); Computer science; Biology; Genetics; Physics; Gene","score_opus":0.020230871594748744,"score_gpt":0.2870178750129633,"score_spread":0.26678700341821454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W178313090","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42621166,0.009396678,0.54698867,0.0068781083,0.00056054373,0.00019121204,0.0016076834,0.0010211994,0.0071443333],"genre_scores_gemma":[0.9064736,0.0008875971,0.09000171,0.0004154028,0.00024008629,0.00013619539,0.0012534506,0.0003182147,0.00027382132],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9245274,0.048760112,0.005235866,0.009884989,0.010389932,0.001201658],"domain_scores_gemma":[0.39228258,0.5272963,0.020413386,0.033259887,0.022585794,0.004162123],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.072060265,0.0012371134,0.0031091385,0.009774715,0.003480337,0.0070874062,0.004414887,0.007941742,0.002198511],"category_scores_gemma":[0.5455394,0.0014785914,0.0023499127,0.0108625395,0.00696162,0.018161373,0.0052485717,0.005686683,0.0008775573],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016811043,0.0002904635,0.31672585,0.003020161,0.008440174,0.0016650537,0.015435068,0.1949884,0.004707346,0.08694792,0.015769035,0.3503295],"study_design_scores_gemma":[0.0001478101,0.000353695,0.04204838,0.00066763797,0.0012815502,0.0010781874,0.004189245,0.25079006,0.003043693,0.69127864,0.0048469114,0.00027417648],"about_ca_topic_score_codex":0.0020657287,"about_ca_topic_score_gemma":0.0029377271,"teacher_disagreement_score":0.072060265,"about_ca_system_score_codex":0.0017759706,"about_ca_system_score_gemma":0.0015243406,"threshold_uncertainty_score":0.38109565},"labels":[],"label_agreement":null},{"id":"W1794390359","doi":"10.1007/978-3-540-69828-9_17","title":"Chemical Knowledge for the Semantic Web","year":2008,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":28,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Ontology; Semantic Web; PubChem; Knowledge representation and reasoning; DrugBank; Information retrieval; Semantic Web Rule Language; Web Ontology Language; Sublanguage; OWL-S; World Wide Web; Social Semantic Web; Natural language processing; Artificial intelligence; Semantic analytics","score_opus":0.022038828262771693,"score_gpt":0.2728477746343452,"score_spread":0.25080894637157347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1794390359","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007091969,0.018084519,0.81416285,0.018810455,0.002112288,0.00023709433,0.0058111944,0.0075895796,0.12610006],"genre_scores_gemma":[0.12211383,0.024118857,0.79414654,0.0035232943,0.0013741426,0.00033712393,0.013830783,0.0013989792,0.039156433],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990957,0.00017468525,0.00013071374,0.00014170005,0.00040351463,0.00005373708],"domain_scores_gemma":[0.99897707,0.0003555176,0.00004670105,0.00040635743,0.00014890432,0.00006549126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015727205,0.0007347766,0.0007990908,0.0034448104,0.0010871378,0.0046713552,0.0016152199,0.0019614252,0.013221037],"category_scores_gemma":[0.003311269,0.0007440449,0.0012949032,0.0033111463,0.0018171647,0.015232907,0.0037879704,0.0023440577,0.005959224],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035579014,0.000043046995,0.00012574162,0.00037975094,0.000036211575,0.00016494568,0.00015338248,0.0024091236,0.0014101852,0.8339438,0.031123312,0.130175],"study_design_scores_gemma":[0.000008639764,0.00000530864,0.000079320474,0.00013016195,0.000019223255,0.0001160243,0.000050316652,0.009348988,0.0010570649,0.80588055,0.18329366,0.000010766109],"about_ca_topic_score_codex":0.0022699004,"about_ca_topic_score_gemma":0.0028383392,"teacher_disagreement_score":0.013221037,"about_ca_system_score_codex":0.0012921442,"about_ca_system_score_gemma":0.0014027802,"threshold_uncertainty_score":0.044228733},"labels":[],"label_agreement":null},{"id":"W1858261188","doi":"10.1007/978-3-540-73599-1_60","title":"Semantic Web Framework for Knowledge-Centric Clinical Decision Support Systems","year":2007,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Semantic Web; Social Semantic Web; Clinical decision support system; Decision support system; World Wide Web; Knowledge management; Artificial intelligence","score_opus":0.04553683976924533,"score_gpt":0.3575536375243903,"score_spread":0.31201679775514496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1858261188","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002235879,0.0012427793,0.9813131,0.0014984034,0.00020293989,0.00026757963,0.0018726006,0.0057090675,0.005657637],"genre_scores_gemma":[0.06477798,0.002054771,0.919388,0.0010547037,0.00022862785,0.0005445486,0.0071976664,0.0005594808,0.0041942657],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9973314,0.00068170036,0.00056095945,0.000352929,0.0008896978,0.00018342293],"domain_scores_gemma":[0.9984725,0.00060608104,0.00011300241,0.00030259017,0.0003433907,0.00016242177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00424445,0.00091549475,0.0017646243,0.0040957434,0.0014310313,0.0072669145,0.0029395348,0.0022256167,0.0034214908],"category_scores_gemma":[0.0037075242,0.0007398747,0.0024014362,0.0043995865,0.0013222849,0.0056590256,0.0036483903,0.0023227492,0.0020529844],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033333406,0.00032386387,0.0008948072,0.001170562,0.0004461771,0.0011098873,0.00071764307,0.03899033,0.0051011858,0.6899467,0.03618514,0.22478032],"study_design_scores_gemma":[0.00010132829,0.000063637104,0.0003852168,0.00045658584,0.0003212443,0.0006142324,0.00023458397,0.2015729,0.006945772,0.5912452,0.19798496,0.00007428708],"about_ca_topic_score_codex":0.0056144805,"about_ca_topic_score_gemma":0.0059673022,"teacher_disagreement_score":0.0072669145,"about_ca_system_score_codex":0.00174484,"about_ca_system_score_gemma":0.0037232314,"threshold_uncertainty_score":0.02244705},"labels":[],"label_agreement":null},{"id":"W186727283","doi":"","title":"Exploratory Reverse Mapping of ICD-10-CA to SNOMED CT.","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Matching (statistics); Computer science; Medicine; Information retrieval; Terminology; Pathology","score_opus":0.03728912500705173,"score_gpt":0.2624175818137827,"score_spread":0.225128456806731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W186727283","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3493703,0.0019371178,0.5099379,0.0025946663,0.00051048206,0.005082473,0.07546841,0.008033584,0.04706514],"genre_scores_gemma":[0.43107525,0.00071671925,0.50339264,0.00046511155,0.000040666953,0.0011782409,0.05652087,0.0009946837,0.0056158323],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99505794,0.0014483619,0.0003914202,0.0006263592,0.0022846682,0.00019122033],"domain_scores_gemma":[0.966964,0.01774395,0.001711779,0.0030801713,0.010257044,0.00024304597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073154955,0.0006394358,0.00036555238,0.006735346,0.0013989813,0.0019690804,0.0012369184,0.00035610932,0.0031957638],"category_scores_gemma":[0.045474764,0.000279976,0.0008626567,0.007301264,0.0008102684,0.00094699906,0.002022955,0.0010584636,0.0007499104],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010434872,0.0004098825,0.10383305,0.005356797,0.0008822689,0.0039770543,0.03539058,0.017756023,0.049450085,0.0319086,0.054675803,0.6953164],"study_design_scores_gemma":[0.00014188752,0.0003405887,0.16997154,0.0022690275,0.0010078739,0.0072070975,0.02822772,0.10130442,0.0908092,0.032548163,0.5657008,0.0004716203],"about_ca_topic_score_codex":0.22044171,"about_ca_topic_score_gemma":0.27449232,"teacher_disagreement_score":0.22044171,"about_ca_system_score_codex":0.003258984,"about_ca_system_score_gemma":0.010173172,"threshold_uncertainty_score":0.43831718},"labels":[],"label_agreement":null},{"id":"W1868020197","doi":"10.1111/coin.12009","title":"EXPLORING A SUBGRAPH MATCHING APPROACH FOR EXTRACTING BIOLOGICAL EVENTS FROM LITERATURE","year":2013,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Parsing; Subgraph isomorphism problem; Event (particle physics); Task (project management); Matching (statistics); Artificial intelligence; Natural language processing; Data mining; Graph; Machine learning; Information retrieval; Theoretical computer science","score_opus":0.15139652313491275,"score_gpt":0.31821172745195253,"score_spread":0.16681520431703978,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1868020197","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.064576045,0.0026127747,0.9033359,0.0010758109,0.0001019194,0.0009951477,0.01142266,0.009765772,0.006113867],"genre_scores_gemma":[0.1875591,0.0013200461,0.77946305,0.00021432972,0.00007157872,0.0004230184,0.028542556,0.00036664223,0.0020397345],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987006,0.0003214294,0.00017032039,0.0004348709,0.0002979417,0.00007466949],"domain_scores_gemma":[0.9974915,0.0015748166,0.00022889904,0.000285538,0.00034847643,0.00007074622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001296755,0.000977404,0.0006956208,0.014043634,0.0009963424,0.0010840766,0.0012452062,0.0011277609,0.003034912],"category_scores_gemma":[0.0059992266,0.0004797214,0.002141124,0.008658275,0.000625974,0.0024781865,0.0014936371,0.00067225064,0.001151313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039124335,0.0004884224,0.011707447,0.0028880902,0.00071707496,0.0021157707,0.001436997,0.046947695,0.04503703,0.04110897,0.025933672,0.82122755],"study_design_scores_gemma":[0.00019672635,0.00030986316,0.0141187,0.00040104438,0.0008878597,0.0021992316,0.0013993378,0.6019555,0.03834428,0.25033158,0.089722484,0.00013340871],"about_ca_topic_score_codex":0.008936291,"about_ca_topic_score_gemma":0.016669065,"teacher_disagreement_score":0.014043634,"about_ca_system_score_codex":0.001008406,"about_ca_system_score_gemma":0.0022675293,"threshold_uncertainty_score":0.017768562},"labels":[],"label_agreement":null},{"id":"W1888011339","doi":"","title":"Clinical Information Retrieval using Document and PICO Structure","year":2010,"lang":"en","type":"article","venue":"North American Chapter of the Association for Computational Linguistics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"","keywords":"Information retrieval; Weighting; Computer science; Document Structure Description; Data mining; Medicine; XML; World Wide Web","score_opus":0.011648740625397166,"score_gpt":0.30021825393840634,"score_spread":0.2885695133130092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1888011339","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09959158,0.008047386,0.86512655,0.0025978282,0.00042589995,0.0011134741,0.008850271,0.0063504893,0.007896395],"genre_scores_gemma":[0.43157056,0.002981223,0.5437069,0.00061109726,0.00067792006,0.0011779866,0.014505244,0.0004238821,0.004345237],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9957455,0.0013592246,0.0006999595,0.00078530255,0.0012528996,0.00015712013],"domain_scores_gemma":[0.9880932,0.0070016226,0.0013264068,0.0014231586,0.0018214962,0.00033406614],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003358747,0.0009640335,0.0017973335,0.016395556,0.0011323818,0.0029326128,0.0010734331,0.0013351853,0.0026124376],"category_scores_gemma":[0.024944788,0.00062015717,0.0014417266,0.012329265,0.0008873081,0.007831783,0.0028730142,0.0010434838,0.0016897998],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001067427,0.00033840025,0.011994122,0.0015615044,0.0003679744,0.0004783221,0.0013727028,0.028672004,0.018120792,0.036148224,0.022620823,0.87725776],"study_design_scores_gemma":[0.00036275154,0.0009802581,0.012623066,0.00043795136,0.00054618064,0.0019875797,0.0010121303,0.7502884,0.01876292,0.15061189,0.06213248,0.00025430674],"about_ca_topic_score_codex":0.0048958743,"about_ca_topic_score_gemma":0.004865097,"teacher_disagreement_score":0.016395556,"about_ca_system_score_codex":0.0017848689,"about_ca_system_score_gemma":0.0022586165,"threshold_uncertainty_score":0.017762959},"labels":[],"label_agreement":null},{"id":"W188885233","doi":"","title":"Situational Modeling: Defining Molecular Roles in Biochemical Pathways and Reactions.","year":2008,"lang":"en","type":"article","venue":"Research Publications (Maastricht University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Representation (politics); Computational biology; Computer science; Situational ethics; Cognitive science; Data science; Biology; Psychology; Social psychology; Political science","score_opus":0.08640423311333267,"score_gpt":0.3037219930455195,"score_spread":0.21731775993218683,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W188885233","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013863467,0.000715871,0.96755755,0.0010537128,0.00015549149,0.00019273764,0.0029881939,0.0048886156,0.008584331],"genre_scores_gemma":[0.24609244,0.0010548657,0.7439358,0.00029923965,0.00007169509,0.00028828098,0.0049254405,0.0003352041,0.00299706],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9990181,0.0003898701,0.00015442164,0.0001534182,0.00022772902,0.00005643439],"domain_scores_gemma":[0.9985372,0.0007338044,0.0001581779,0.0003053519,0.00014845586,0.00011706958],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002198466,0.0008950607,0.00038795275,0.0023080134,0.00080778374,0.0030204873,0.0013625557,0.0009579299,0.0033762804],"category_scores_gemma":[0.0039795,0.00046146638,0.0016743778,0.0015247798,0.0010924195,0.0050089434,0.002006651,0.0012209709,0.0008869371],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002743831,0.00014179069,0.0061007617,0.0010190735,0.00016315214,0.0013900598,0.0024004364,0.08637345,0.007613062,0.7383585,0.012730167,0.14343514],"study_design_scores_gemma":[0.000060979193,0.000059518236,0.0012605732,0.0004264993,0.00024740255,0.0008777257,0.0009805264,0.4404105,0.010486962,0.3791737,0.16596544,0.000050106093],"about_ca_topic_score_codex":0.0066194176,"about_ca_topic_score_gemma":0.010805602,"teacher_disagreement_score":0.0066194176,"about_ca_system_score_codex":0.00094756536,"about_ca_system_score_gemma":0.0018705291,"threshold_uncertainty_score":0.013161778},"labels":[],"label_agreement":null},{"id":"W188993809","doi":"","title":"Use of OWL 2 to Facilitate a Biomedical Knowledge Base Extracted from the GENIA Corpus.","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Annotation; Knowledge base; Taxonomy (biology); Artificial intelligence; Natural language processing; Domain (mathematical analysis); Generalization; Set (abstract data type); Process (computing); Fuzzy logic; Hierarchy; Information retrieval; Programming language","score_opus":0.12109027724039965,"score_gpt":0.2819199807081979,"score_spread":0.16082970346779824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W188993809","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09619704,0.0026957917,0.62040216,0.008780073,0.0009962333,0.0025858306,0.17958178,0.03110223,0.05765894],"genre_scores_gemma":[0.1591969,0.0010754623,0.6939501,0.0013370032,0.00010782552,0.0012391263,0.13333891,0.0017078716,0.008046737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992644,0.00022952071,0.00014746143,0.00012757281,0.00019965538,0.000031415217],"domain_scores_gemma":[0.9944299,0.003435987,0.00035303453,0.0005980205,0.001028135,0.00015491564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022080285,0.0004124728,0.00040294282,0.004521316,0.00094329164,0.0019822225,0.00075844885,0.00051539415,0.005714262],"category_scores_gemma":[0.010923486,0.0004063809,0.00055781787,0.0027594564,0.00048307533,0.002383518,0.0024892609,0.0011567788,0.0020232035],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055414735,0.00042675322,0.011165075,0.005278609,0.0003119021,0.007937408,0.008441644,0.012345081,0.047146067,0.07971683,0.21402863,0.61264783],"study_design_scores_gemma":[0.0001672644,0.00009364822,0.01502711,0.001227032,0.00027933167,0.0034295514,0.0024211842,0.068948455,0.035480354,0.040433303,0.83237195,0.000120726865],"about_ca_topic_score_codex":0.012113639,"about_ca_topic_score_gemma":0.018299447,"teacher_disagreement_score":0.012113639,"about_ca_system_score_codex":0.0010931988,"about_ca_system_score_gemma":0.002634503,"threshold_uncertainty_score":0.024086237},"labels":[],"label_agreement":null},{"id":"W1901558184","doi":"","title":"Analysis of polarity information in medical text.","year":2005,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":83,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Bigram; Computer science; Artificial intelligence; Natural language processing; Outcome (game theory); Sentence; Negation; Polarity (international relations); Natural language; Context (archaeology); Knowledge base; Generalization; Feature (linguistics); Machine learning; Question answering; Linguistics; Mathematics; Trigram","score_opus":0.01117752964617533,"score_gpt":0.24860632441329938,"score_spread":0.23742879476712406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1901558184","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5640345,0.031859547,0.23057961,0.014800192,0.0037014312,0.001876448,0.08294895,0.0028148359,0.06738441],"genre_scores_gemma":[0.83025116,0.006560896,0.12453399,0.0011632688,0.002216297,0.00065337133,0.03173146,0.000129976,0.002759551],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99815625,0.00069611747,0.0003016068,0.00022176106,0.000555249,0.00006897751],"domain_scores_gemma":[0.9752736,0.019332813,0.0024590436,0.0004752847,0.0022301334,0.00022910263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001927651,0.00046473768,0.0004037757,0.0075363982,0.00057741004,0.0012934986,0.00030496548,0.0005874786,0.0034564869],"category_scores_gemma":[0.019014949,0.00013529908,0.0003410459,0.0037020345,0.0005537059,0.0025026596,0.0007902642,0.00051614235,0.001360593],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017412428,0.0002689255,0.046910975,0.0067053596,0.00030831463,0.003604345,0.0027164451,0.0024686244,0.07563733,0.023422685,0.046268605,0.7899471],"study_design_scores_gemma":[0.00032223237,0.0013809106,0.23982885,0.0033952035,0.0015252945,0.018698106,0.00918002,0.11339516,0.09401779,0.23271602,0.2851394,0.00040099942],"about_ca_topic_score_codex":0.0003544529,"about_ca_topic_score_gemma":0.000562083,"teacher_disagreement_score":0.0075363982,"about_ca_system_score_codex":0.00039647653,"about_ca_system_score_gemma":0.0005645819,"threshold_uncertainty_score":0.011563063},"labels":[],"label_agreement":null},{"id":"W1901898645","doi":"10.5489/cuaj.795","title":"Quantifying CUA’s progress","year":2013,"lang":"en","type":"article","venue":"Canadian Urological Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Canadian Urological Association","funders":"","keywords":"Computer science","score_opus":0.020866738150760046,"score_gpt":0.2555186236013835,"score_spread":0.23465188545062346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1901898645","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.77726513,0.00786928,0.078224644,0.0069165085,0.0014158653,0.0003895955,0.027709626,0.0069538467,0.09325553],"genre_scores_gemma":[0.94622254,0.0009362168,0.036988966,0.00018076209,0.00014774491,0.00011962568,0.007923303,0.0003464595,0.0071344348],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9940925,0.0011247866,0.0005219671,0.0015306479,0.0020852515,0.00064479135],"domain_scores_gemma":[0.96447664,0.014025718,0.003277541,0.0036196425,0.012797335,0.0018029974],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072975187,0.00062884693,0.0007661547,0.012545394,0.0021901014,0.0057333503,0.0010595083,0.0012712404,0.0038347617],"category_scores_gemma":[0.043805473,0.00027191444,0.0009088927,0.01032765,0.0008481656,0.0061119352,0.0027499658,0.0012274834,0.0017265552],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060577196,0.00025479676,0.46481627,0.0007087783,0.00042348943,0.0004370508,0.00437296,0.020852273,0.005876228,0.032255586,0.02825628,0.4411406],"study_design_scores_gemma":[0.00004547022,0.00029230112,0.3996655,0.00051748723,0.00071648246,0.0012170987,0.012973202,0.285356,0.024395827,0.049264032,0.22525157,0.00030508783],"about_ca_topic_score_codex":0.05380596,"about_ca_topic_score_gemma":0.04656913,"teacher_disagreement_score":0.05380596,"about_ca_system_score_codex":0.0032122598,"about_ca_system_score_gemma":0.0048720734,"threshold_uncertainty_score":0.10698551},"labels":[],"label_agreement":null},{"id":"W190884861","doi":"10.1007/978-3-642-21043-3_17","title":"Extracting Relations between Diseases, Treatments, and Tests from Clinical Data","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Artificial intelligence; Relation (database); Sentence; Support vector machine; Identification (biology); Natural language processing; Task (project management); Set (abstract data type); Projection (relational algebra); Data set; Random projection; Relationship extraction; Pattern recognition (psychology); Data mining; Algorithm","score_opus":0.08944012625585655,"score_gpt":0.34846279627255444,"score_spread":0.2590226700166979,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W190884861","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17597787,0.044972744,0.52730423,0.010364008,0.0011364635,0.0019457975,0.21616332,0.0068896576,0.0152459275],"genre_scores_gemma":[0.2289993,0.015711468,0.5365824,0.0014790279,0.0007037334,0.00068658363,0.21293618,0.00039744526,0.0025038354],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978255,0.0004387941,0.00051285647,0.0005446167,0.00057514984,0.00010307474],"domain_scores_gemma":[0.98835665,0.009336966,0.0008581529,0.00072087476,0.0004940475,0.00023330841],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002888511,0.0014832378,0.0018641371,0.011674317,0.0007172002,0.002802415,0.0015308493,0.0012543573,0.0038339656],"category_scores_gemma":[0.013566762,0.00067171844,0.0032639727,0.010027075,0.0005806114,0.0024017459,0.0018548355,0.0015593858,0.0025085374],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090558565,0.000635967,0.07718782,0.006802591,0.0014853938,0.004196986,0.0013389962,0.009522604,0.020742502,0.012295531,0.035637803,0.8292483],"study_design_scores_gemma":[0.0007139538,0.0011741936,0.17381234,0.005683246,0.009575779,0.017923893,0.004347801,0.13734199,0.051821947,0.32385147,0.27329174,0.00046162424],"about_ca_topic_score_codex":0.003209531,"about_ca_topic_score_gemma":0.0053727184,"teacher_disagreement_score":0.011674317,"about_ca_system_score_codex":0.00074269366,"about_ca_system_score_gemma":0.003018562,"threshold_uncertainty_score":0.015276134},"labels":[],"label_agreement":null},{"id":"W1910277422","doi":"10.2196/medinform.4450","title":"Technology for Large-Scale Translation of Clinical Practice Guidelines: A Pilot Study of the Performance of a Hybrid Human and Computer-Assisted Approach","year":2015,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Terminology; Computer science; Computer-assisted translation; Machine translation; Process (computing); Scale (ratio); Translation (biology); Knowledge translation; Quality (philosophy); Natural language processing; World Wide Web; Artificial intelligence; Data science; Knowledge management; Linguistics; Programming language","score_opus":0.11594430204313633,"score_gpt":0.4130645175378309,"score_spread":0.2971202154946946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1910277422","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.993572,0.000057645455,0.00317276,0.000044207707,0.000023297645,0.0024585489,0.000069589485,0.0000646318,0.0005372195],"genre_scores_gemma":[0.9636803,0.00014166479,0.028139733,0.0001824688,0.00007476031,0.0069376337,0.0002166668,0.000039585502,0.0005870489],"study_design_codex":"design_other","study_design_gemma":"nonrandomized_trial","domain_scores_codex":[0.9732159,0.020524133,0.001842969,0.001992644,0.0017917939,0.00063272053],"domain_scores_gemma":[0.9017496,0.08108145,0.004075002,0.0063490337,0.003999591,0.0027452698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028022315,0.001091319,0.0011835244,0.0010629912,0.00073032686,0.0011993371,0.0012861443,0.0011939774,0.004019322],"category_scores_gemma":[0.068121426,0.00082200556,0.0009617125,0.0007389502,0.0012314644,0.0017337062,0.0019609595,0.0009176264,0.0010119115],"study_design_candidate":"nonrandomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.14524491,0.18600427,0.050731737,0.0030606932,0.0015915985,0.0008558876,0.026727913,0.008378393,0.02736516,0.000612603,0.001863114,0.5475637],"study_design_scores_gemma":[0.047137175,0.7841795,0.11251359,0.0003074385,0.0015958439,0.0008108024,0.0059938794,0.024222095,0.01630218,0.0008665249,0.0056771594,0.00039391418],"about_ca_topic_score_codex":0.0009071218,"about_ca_topic_score_gemma":0.000827229,"teacher_disagreement_score":0.028022315,"about_ca_system_score_codex":0.00075642776,"about_ca_system_score_gemma":0.001696437,"threshold_uncertainty_score":0.14819795},"labels":[],"label_agreement":null},{"id":"W1913960352","doi":"10.1111/j.1945-1474.2011.00174.x","title":"An Integrated Methodology for Process Improvement and Delivery System Visualization at a Multidisciplinary Cancer Center","year":2011,"lang":"en","type":"article","venue":"Journal for Healthcare Quality","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"CARE Canada","funders":"","keywords":"Multidisciplinary approach; Workflow; Computer science; Process (computing); Data collection; Process management; Visualization; Knowledge management; Systems engineering; Engineering; Data mining; Database","score_opus":0.21856003546988545,"score_gpt":0.4892988647211996,"score_spread":0.27073882925131415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1913960352","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005584737,0.0000665671,0.99083745,0.00027539456,0.000018477227,0.00054724084,0.00023734111,0.001293608,0.0011392183],"genre_scores_gemma":[0.02400168,0.000059799568,0.97493285,0.000023020602,0.000003793554,0.00038867246,0.00028611583,0.000040890365,0.00026316027],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9909763,0.003903192,0.0013068422,0.001216724,0.0023079624,0.00028902138],"domain_scores_gemma":[0.9888777,0.0056317896,0.0011348283,0.0013602286,0.0027594469,0.00023598918],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01081448,0.001277764,0.0007923213,0.0060880766,0.0012757751,0.004734807,0.0017271512,0.0011526116,0.0026259157],"category_scores_gemma":[0.015103147,0.0006825878,0.0017507293,0.004844432,0.0010479864,0.003974738,0.0024467295,0.0011014191,0.0005306777],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002355505,0.0007023057,0.010436717,0.0024119231,0.00036216993,0.0009559039,0.010844219,0.077902064,0.015924433,0.11028601,0.0062231338,0.7637156],"study_design_scores_gemma":[0.0002660562,0.0007909281,0.00762731,0.0013398236,0.00046986056,0.0011852115,0.0073272185,0.75816697,0.03097003,0.09152231,0.10005586,0.00027844904],"about_ca_topic_score_codex":0.00559698,"about_ca_topic_score_gemma":0.0055996585,"teacher_disagreement_score":0.01081448,"about_ca_system_score_codex":0.002331292,"about_ca_system_score_gemma":0.0069297003,"threshold_uncertainty_score":0.0571931},"labels":[],"label_agreement":null},{"id":"W191751253","doi":"","title":"The OWL of Biomedical Investigations.","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"","keywords":"Computer science; Ontology; OWL-S; Web Ontology Language; Data science; Process (computing); Open Biomedical Ontologies; World Wide Web; Semantic Web; Programming language; Social Semantic Web","score_opus":0.02349250003918839,"score_gpt":0.2602904352441354,"score_spread":0.23679793520494702,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W191751253","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004498746,0.0045368825,0.69272786,0.01219598,0.0015059499,0.0021180555,0.14447887,0.042357355,0.09558035],"genre_scores_gemma":[0.07867291,0.007381244,0.6869522,0.01166445,0.00082238606,0.002268566,0.18099247,0.0040583145,0.027187483],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9971585,0.00075816666,0.00063059863,0.00034660884,0.0009600174,0.00014606945],"domain_scores_gemma":[0.9949685,0.0025244097,0.0005095763,0.0008592977,0.00067170896,0.00046653755],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004454489,0.00093956874,0.0008613146,0.005605463,0.0017122965,0.0043153004,0.0024564036,0.0014739296,0.01090654],"category_scores_gemma":[0.011081035,0.00084777985,0.0013995776,0.004647178,0.0011737003,0.005759247,0.0045321006,0.0023899043,0.005737633],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022343236,0.00019948152,0.0039613643,0.0043333946,0.000256939,0.0015693184,0.0014883862,0.0034018946,0.0044800597,0.36126328,0.37644345,0.24237904],"study_design_scores_gemma":[0.000093412586,0.000028787003,0.0023421678,0.0005268033,0.00007550029,0.0011254243,0.00043408485,0.009813116,0.0019627719,0.13273534,0.85080326,0.000059309474],"about_ca_topic_score_codex":0.018713031,"about_ca_topic_score_gemma":0.02270467,"teacher_disagreement_score":0.018713031,"about_ca_system_score_codex":0.002110495,"about_ca_system_score_gemma":0.0054473705,"threshold_uncertainty_score":0.0372082},"labels":[],"label_agreement":null},{"id":"W1926090033","doi":"10.1186/1471-2105-4-51","title":"A knowledge discovery object model API for Java","year":2003,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; BC Cancer Agency","funders":"Genome Canada","keywords":"Computer science; Application programming interface; Java; Documentation; Ontology; Software engineering; Software; Software mining; Interface (matter); Database; Software development; Programming language; Software construction; Operating system","score_opus":0.03862548388309531,"score_gpt":0.2979191660009343,"score_spread":0.259293682117839,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1926090033","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009842096,0.00030588254,0.751685,0.00054373284,0.00013258368,0.00060301414,0.009898438,0.22465503,0.011192108],"genre_scores_gemma":[0.021866677,0.001074994,0.82830244,0.0016790422,0.00017014405,0.0030655721,0.061599288,0.06367046,0.018571483],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99764687,0.00028876893,0.00044601163,0.000431862,0.0009687356,0.00021777311],"domain_scores_gemma":[0.9957535,0.001989078,0.00044955843,0.000703116,0.00084003276,0.00026464535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004545225,0.0017282063,0.0013721981,0.0028877475,0.0009928542,0.004543394,0.0049923644,0.0021087204,0.02293569],"category_scores_gemma":[0.008284583,0.0020479127,0.0029174776,0.0024809723,0.00077951717,0.005371934,0.0037195135,0.0038583335,0.02465017],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011532914,0.000779229,0.004713144,0.0026623285,0.00034574282,0.0010803294,0.000925762,0.005808966,0.018203529,0.09131889,0.42824087,0.44476795],"study_design_scores_gemma":[0.00075218943,0.00013403773,0.0028783542,0.0005621311,0.00017431406,0.0014816893,0.00013880018,0.067423806,0.018319096,0.08270475,0.8251466,0.00028421017],"about_ca_topic_score_codex":0.004103476,"about_ca_topic_score_gemma":0.0035423788,"teacher_disagreement_score":0.02293569,"about_ca_system_score_codex":0.0014132801,"about_ca_system_score_gemma":0.0028145702,"threshold_uncertainty_score":0.07672751},"labels":[],"label_agreement":null},{"id":"W192939834","doi":"10.4137/bbi.s451","title":"Ontologies for Bioinformatics","year":2008,"lang":"en","type":"article","venue":"Bioinformatics and Biology Insights","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Interoperability; Computer science; Ontology; Data science; Context (archaeology); Semantic interoperability; Open Biomedical Ontologies; Semantics (computer science); Semantic Web; Meaning (existential); IDEF5; Semantic integration; World Wide Web; Knowledge management; Upper ontology; Semantic Web Stack; Ontology alignment; Biology","score_opus":0.05290572275187795,"score_gpt":0.2885351413558813,"score_spread":0.23562941860400338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W192939834","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015778623,0.07355103,0.74164146,0.046099246,0.0046557016,0.0005709061,0.003785754,0.0058191996,0.12229882],"genre_scores_gemma":[0.04762332,0.065267086,0.8346211,0.012246134,0.0047982447,0.0011426251,0.0088894,0.0011221432,0.024289927],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99168515,0.0036731367,0.001052437,0.0010872002,0.0021044756,0.00039754462],"domain_scores_gemma":[0.98967135,0.0051363544,0.0007007119,0.0023190985,0.0015579449,0.0006146054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0086817825,0.0015564805,0.0019467329,0.0061539076,0.0036455418,0.0110877715,0.004062039,0.005284216,0.014646111],"category_scores_gemma":[0.020355932,0.0009603506,0.0023754823,0.009236624,0.006362775,0.016648864,0.0067832856,0.0069039688,0.010127194],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001928399,0.000024016437,0.000168425,0.0007323689,0.000048188253,0.0001484886,0.00048708927,0.001527874,0.00030000045,0.864491,0.0440271,0.088026136],"study_design_scores_gemma":[0.000007647915,0.0000045794086,0.00006959586,0.0004041079,0.000011982927,0.00013232628,0.00013736865,0.0019069093,0.00010021182,0.730413,0.26679662,0.000015590233],"about_ca_topic_score_codex":0.005404662,"about_ca_topic_score_gemma":0.0037330214,"teacher_disagreement_score":0.014646111,"about_ca_system_score_codex":0.0039603454,"about_ca_system_score_gemma":0.0071091163,"threshold_uncertainty_score":0.04899615},"labels":[],"label_agreement":null},{"id":"W1931319488","doi":"10.1038/npre.2007.945.1","title":"Bridging the gap between social tagging and semantic annotation: E.D. the Entity Describer","year":2007,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Providence Health Care","funders":"Genome Prairie; Genome Alberta","keywords":"Computer science; Annotation; Bridging (networking); RDF; Information retrieval; Semantic Web; World Wide Web; Artificial intelligence","score_opus":0.02961210909848352,"score_gpt":0.3225281930507134,"score_spread":0.2929160839522299,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1931319488","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013318755,0.0008781291,0.9519235,0.0105493795,0.0002003792,0.00015747291,0.0007443514,0.0033019406,0.018926065],"genre_scores_gemma":[0.23489828,0.0018693174,0.74156934,0.0021076503,0.0003487437,0.00025828936,0.0025336668,0.0019304996,0.014484196],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9824806,0.010127111,0.001198079,0.0020037757,0.003760822,0.0004296724],"domain_scores_gemma":[0.9319052,0.03467353,0.0019601774,0.026408434,0.0041420804,0.0009105905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027081467,0.0007065624,0.0011213372,0.004713783,0.0026693754,0.008005675,0.0029525133,0.00291575,0.0065598697],"category_scores_gemma":[0.048340566,0.0009734538,0.0011431779,0.004802657,0.004886254,0.022941563,0.009905927,0.002859911,0.003215576],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029355593,0.00015685278,0.0025680666,0.0003124389,0.000068207955,0.0005416507,0.00338904,0.004518724,0.0052407742,0.7530576,0.023690019,0.20616305],"study_design_scores_gemma":[0.00009595103,0.00005581402,0.0010178942,0.00034317438,0.0001386095,0.000990889,0.0020136938,0.11951813,0.027870288,0.4954214,0.35238212,0.00015211565],"about_ca_topic_score_codex":0.0096648205,"about_ca_topic_score_gemma":0.005286058,"teacher_disagreement_score":0.027081467,"about_ca_system_score_codex":0.0031275707,"about_ca_system_score_gemma":0.0032747877,"threshold_uncertainty_score":0.14322221},"labels":[],"label_agreement":null},{"id":"W1933893559","doi":"10.1111/j.1365-2753.2003.00432.x","title":"Knowledge base of scientific gnosis: III. Gnostic occurrence relations as regression functions","year":2004,"lang":"en","type":"article","venue":"Journal of Evaluation in Clinical Practice","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Biostatistics; Epidemiology; Library science; Face (sociological concept); Medicine; Sociology; Computer science; Social science; Pathology","score_opus":0.12468352982423304,"score_gpt":0.49446326362534393,"score_spread":0.3697797338011109,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1933893559","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046133418,0.0026138953,0.9217758,0.001899969,0.00018356329,0.00041383237,0.008226474,0.006716378,0.012036652],"genre_scores_gemma":[0.50292563,0.002705579,0.46969146,0.0005620493,0.00027751984,0.00044476858,0.017446464,0.0004202102,0.005526402],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99746954,0.00060963875,0.00042471333,0.00047457844,0.00087935367,0.0001421536],"domain_scores_gemma":[0.99318177,0.0039654416,0.0003352768,0.0010350818,0.0013302075,0.00015222836],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004188385,0.00054245244,0.0009855537,0.0064046923,0.00067699095,0.0035127583,0.0015414585,0.0011363897,0.0049237134],"category_scores_gemma":[0.019263191,0.0004040439,0.0012555084,0.004215526,0.00085521874,0.0041886833,0.0015914683,0.0014190808,0.0022279455],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065235514,0.00038847953,0.023425844,0.0009679966,0.00051858876,0.00096418645,0.0005651751,0.056020766,0.006429304,0.081949085,0.018217472,0.80990076],"study_design_scores_gemma":[0.000104658764,0.00017051419,0.012903551,0.00077514764,0.00094131724,0.001449354,0.0003115907,0.7206523,0.016670225,0.19859913,0.047306277,0.00011601529],"about_ca_topic_score_codex":0.00676562,"about_ca_topic_score_gemma":0.0052861334,"teacher_disagreement_score":0.00676562,"about_ca_system_score_codex":0.0010198676,"about_ca_system_score_gemma":0.0017476835,"threshold_uncertainty_score":0.022150576},"labels":[],"label_agreement":null},{"id":"W1934573562","doi":"10.1038/npre.2010.4270.2","title":"Formulating MEDLINE queries for article retrieval based on PubMed exemplars","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"British Columbia Knowledge Development Fund; National Institutes of Health","keywords":"Computer science; Information retrieval; Search engine indexing; Set (abstract data type); Bigram; Task (project management); Result set; Function (biology); Process (computing); Recall; Natural language processing","score_opus":0.014892616225833525,"score_gpt":0.29648085099582633,"score_spread":0.2815882347699928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1934573562","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12663923,0.003213281,0.7069574,0.0059507843,0.00039932784,0.003926329,0.047528982,0.09200779,0.013376874],"genre_scores_gemma":[0.12772241,0.0010887196,0.83283365,0.00058345357,0.00017841636,0.00096158136,0.03266838,0.0017806554,0.0021827044],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954383,0.0010536843,0.0012458146,0.00067678874,0.0013525705,0.00023290365],"domain_scores_gemma":[0.9804709,0.014344707,0.001169289,0.0009894549,0.0024822059,0.0005434055],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003857875,0.0019336087,0.0023108523,0.011916991,0.0013339403,0.0044674594,0.0018617972,0.0025042754,0.013060763],"category_scores_gemma":[0.031715255,0.0010972196,0.0022263762,0.006463361,0.00075608765,0.006328223,0.002942056,0.0011501138,0.0053527374],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034168933,0.0007931826,0.019090397,0.010376399,0.00065198774,0.004547473,0.004273445,0.018254835,0.09683853,0.045420606,0.13467988,0.6616564],"study_design_scores_gemma":[0.0015091241,0.0011999988,0.014796443,0.001432851,0.00094010757,0.005262696,0.0049531683,0.49648836,0.12191506,0.086901285,0.26408353,0.0005173482],"about_ca_topic_score_codex":0.0035192708,"about_ca_topic_score_gemma":0.0060635107,"teacher_disagreement_score":0.99614215,"about_ca_system_score_codex":0.0017969743,"about_ca_system_score_gemma":0.0018890917,"threshold_uncertainty_score":0.04369253},"labels":[],"label_agreement":null},{"id":"W1944898984","doi":"10.1387/theoria.12697","title":"Complex Underdetermination and the Units of Clinical Translation","year":2015,"lang":"en","type":"article","venue":"THEORIA An International Journal for Theory History and Foundations of Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Canadian Institutes of Health Research; Genome Alberta; McGill University; University of Pittsburgh","keywords":"Underdetermination; Ontology; Biomarker; Quality (philosophy); Epistemology; Translation (biology); Computer science; Data science; Cognitive science; Engineering ethics; Psychology; Philosophy; Philosophy of science; Biology; Engineering","score_opus":0.20168468295213735,"score_gpt":0.4302379736427979,"score_spread":0.22855329069066058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1944898984","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020709014,0.016206505,0.21082889,0.6378218,0.0050524026,0.000663795,0.00038494248,0.0005299581,0.10780272],"genre_scores_gemma":[0.799878,0.0062085367,0.12908949,0.050903715,0.0052062944,0.0022096366,0.0003008862,0.000678943,0.00552451],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.71252054,0.22051683,0.018954728,0.019731436,0.023872139,0.0044043227],"domain_scores_gemma":[0.497865,0.38684136,0.016915431,0.07087857,0.02249732,0.005002286],"candidate_categories":["metaresearch","sts"],"consensus_categories":[],"category_scores_codex":[0.2697447,0.0009996473,0.002589357,0.007924037,0.008582463,0.024395565,0.005221991,0.0098996945,0.0072928467],"category_scores_gemma":[0.3569403,0.0017363301,0.0018936034,0.0052293576,0.13967241,0.04319901,0.021949667,0.014019898,0.0017323534],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058666286,0.000019443636,0.00071337994,0.00027362938,0.000047915142,0.0001036508,0.0073932414,0.00024185552,0.00013595125,0.9644933,0.005223768,0.021295281],"study_design_scores_gemma":[0.00003299733,0.000013538771,0.0003093473,0.00022664948,0.00001904241,0.000062347455,0.0013601634,0.00051557174,0.00017208463,0.979677,0.017587025,0.000024116738],"about_ca_topic_score_codex":0.004817128,"about_ca_topic_score_gemma":0.0019318829,"teacher_disagreement_score":0.9914175,"about_ca_system_score_codex":0.01353961,"about_ca_system_score_gemma":0.018150434,"threshold_uncertainty_score":0.9005348},"labels":[],"label_agreement":null},{"id":"W1964255443","doi":"10.1371/journal.pbio.1002033","title":"Finding Our Way through Phenotypes","year":2015,"lang":"en","type":"article","venue":"PLoS Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":222,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Simon Fraser University","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; National Cancer Institute; National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council; National Science Foundation","keywords":"Phenomics; Biology; Bottleneck; Data science; Genomics; Systematics; Phenotype; Systems biology; Computational biology; Ecology; Computer science; Genome; Genetics; Taxonomy (biology)","score_opus":0.09074615135526791,"score_gpt":0.3319444958987222,"score_spread":0.2411983445434543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964255443","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08013587,0.007299105,0.6694259,0.08599747,0.0025149619,0.00032554546,0.018660977,0.005578357,0.13006179],"genre_scores_gemma":[0.4780569,0.012579515,0.44312575,0.013790061,0.00075714017,0.0005404897,0.01430067,0.004371694,0.03247779],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9960264,0.0014935934,0.000307396,0.0012265841,0.0007883474,0.0001575674],"domain_scores_gemma":[0.98972064,0.004291371,0.00093051384,0.0030914848,0.0014330759,0.000532929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007045968,0.00087429606,0.0008081544,0.0030410776,0.001942595,0.008789867,0.0017059214,0.0014185498,0.013330121],"category_scores_gemma":[0.03184371,0.00048700316,0.0013425619,0.003227581,0.006622014,0.015173579,0.004812764,0.0034787052,0.003923746],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016206839,0.000080393744,0.042618882,0.0009068168,0.00017817922,0.0010174494,0.015849592,0.0019984054,0.004862915,0.5647018,0.059651006,0.3079725],"study_design_scores_gemma":[0.000011341495,0.000039025213,0.013366746,0.00072856323,0.00011584963,0.0009018735,0.0069478205,0.002336653,0.0025387232,0.5515728,0.42134532,0.00009517914],"about_ca_topic_score_codex":0.004454808,"about_ca_topic_score_gemma":0.0043676603,"teacher_disagreement_score":0.013330121,"about_ca_system_score_codex":0.0016272097,"about_ca_system_score_gemma":0.0030398243,"threshold_uncertainty_score":0.04459375},"labels":[],"label_agreement":null},{"id":"W1964979998","doi":"10.1016/j.jbi.2010.04.008","title":"The ACGT Master Ontology and its applications – Towards an ontology-driven cancer research and management system","year":2010,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University","funders":"National Human Genome Research Institute; Alliance for Cancer Gene Therapy","keywords":"Ontology; Computer science; Workflow; Context (archaeology); Software engineering; Interoperability; Knowledge management; World Wide Web; Database","score_opus":0.10450246665768952,"score_gpt":0.42302906689664665,"score_spread":0.3185266002389571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1964979998","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060933228,0.5801445,0.35731578,0.020873204,0.0029837256,0.0003867565,0.0016114684,0.0026653856,0.02792582],"genre_scores_gemma":[0.036207743,0.5503681,0.38593203,0.0061500715,0.0011755553,0.00035415895,0.006556389,0.0003396899,0.012916207],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99873775,0.0001943553,0.00014169968,0.00014927314,0.00070254656,0.00007435668],"domain_scores_gemma":[0.99840325,0.0005510908,0.00016489715,0.00021416912,0.0005559948,0.00011076617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038775436,0.0005613027,0.0010424232,0.0032557193,0.00048804947,0.00254258,0.0024515546,0.0012457465,0.0012540992],"category_scores_gemma":[0.0033386457,0.00039546375,0.0007501517,0.0054195793,0.0013617248,0.0047255447,0.0021682624,0.002033607,0.0013309714],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034792156,0.000073285315,0.00053929404,0.001811176,0.00008494179,0.00014145473,0.0001885787,0.0016865006,0.0026316626,0.062142003,0.030596811,0.9000695],"study_design_scores_gemma":[0.00001547835,0.000026372134,0.0010733302,0.0012615413,0.000097395096,0.00080208137,0.00016407164,0.0043455902,0.0022426695,0.038516082,0.95141596,0.00003930613],"about_ca_topic_score_codex":0.0068504675,"about_ca_topic_score_gemma":0.006734009,"teacher_disagreement_score":0.0068504675,"about_ca_system_score_codex":0.0028911836,"about_ca_system_score_gemma":0.0065583535,"threshold_uncertainty_score":0.02097708},"labels":[],"label_agreement":null},{"id":"W1965360431","doi":"10.3138/jvme.34.4.510","title":"OLIVER: An Online Library of Images for Veterinary Education and Research","year":2007,"lang":"en","type":"article","venue":"Journal of Veterinary Medical Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Flexibility (engineering); World Wide Web; Multimedia; Search engine indexing","score_opus":0.13496442824090754,"score_gpt":0.46929102494303326,"score_spread":0.3343265967021257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965360431","genre_codex":"other","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054869163,0.011642316,0.08542158,0.004578077,0.0027952027,0.0020464235,0.10342511,0.15121436,0.63338995],"genre_scores_gemma":[0.021533161,0.014528936,0.20259856,0.002881997,0.0022775831,0.001456236,0.10086948,0.033048864,0.62080526],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992454,0.00008452499,0.00007022045,0.00009819059,0.00043261418,0.00006907715],"domain_scores_gemma":[0.9958526,0.00095650845,0.00039472745,0.00074514974,0.001115095,0.00093590986],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010415129,0.0012010576,0.0010201897,0.00894316,0.0011333461,0.0047127027,0.0021066803,0.0013896377,0.35973108],"category_scores_gemma":[0.0050735124,0.0007446138,0.00058941776,0.0066960375,0.00062511663,0.0055043753,0.0035577845,0.0014235243,0.18251544],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015355581,0.00007827941,0.00020394435,0.00092375977,0.000011065448,0.00023378788,0.00017640101,0.00011231095,0.0028415904,0.0021674552,0.71647286,0.27662504],"study_design_scores_gemma":[0.000028702787,0.000034154808,0.001464511,0.00027857488,0.000012110923,0.00041713688,0.00011382874,0.00030191932,0.0011314749,0.0016028249,0.9945815,0.00003331465],"about_ca_topic_score_codex":0.0021382968,"about_ca_topic_score_gemma":0.0065883547,"teacher_disagreement_score":0.9952873,"about_ca_system_score_codex":0.0010034398,"about_ca_system_score_gemma":0.0024630341,"threshold_uncertainty_score":0.9132659},"labels":[],"label_agreement":null},{"id":"W1965466476","doi":"10.1016/j.jbi.2012.05.006","title":"A study of terminology auditors’ performance for UMLS semantic type assignments","year":2012,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"New York Institute of Technology","funders":"U.S. National Library of Medicine","keywords":"Unified Medical Language System; Computer science; Audit; Terminology; Recall; Task (project management); Reliability (semiconductor); Information retrieval; Natural language processing; Sample (material); Measure (data warehouse); Semantics (computer science); Artificial intelligence; Data mining; Accounting; Linguistics","score_opus":0.0362487039255609,"score_gpt":0.3222454193852722,"score_spread":0.2859967154597113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1965466476","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9875542,0.00036125703,0.0076454747,0.00043075162,0.00013298093,0.00015024102,0.0006253134,0.0016142146,0.0014856213],"genre_scores_gemma":[0.9798483,0.00016390756,0.015705757,0.0001254193,0.000045163993,0.00008085293,0.0016878806,0.00042452358,0.0019181248],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9308395,0.030565584,0.00997436,0.008484074,0.01726319,0.0028732128],"domain_scores_gemma":[0.42155084,0.42279577,0.040976703,0.046913978,0.061968748,0.0057939724],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.05288736,0.00075131026,0.0011389599,0.0058322004,0.0023717585,0.004222052,0.0018362332,0.0017182579,0.0018675277],"category_scores_gemma":[0.31002244,0.0007014902,0.0011612583,0.0051192055,0.0014051026,0.004768519,0.002990863,0.002423021,0.0011669362],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008846403,0.003193866,0.5024312,0.00095793867,0.0011008733,0.0011743427,0.01376219,0.014616035,0.035170507,0.0024804624,0.011464429,0.4048017],"study_design_scores_gemma":[0.00081152085,0.0069772513,0.48361295,0.0006847469,0.0018917486,0.004665181,0.022177823,0.32506737,0.118879825,0.0047619008,0.029855324,0.00061435596],"about_ca_topic_score_codex":0.007570338,"about_ca_topic_score_gemma":0.0056326333,"teacher_disagreement_score":0.9471126,"about_ca_system_score_codex":0.0022062925,"about_ca_system_score_gemma":0.005066078,"threshold_uncertainty_score":0.2796985},"labels":[],"label_agreement":null},{"id":"W1966038182","doi":"10.1080/01421590120091087","title":"An analysis, using concept mapping, of diabetic patients' knowledge, before and after patient education","year":2002,"lang":"en","type":"article","venue":"Medical Teacher","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Concept map; Cognition; Cognitive map; Psychology; Medicine; Medical education; Mathematics education; Psychiatry","score_opus":0.01091150840097225,"score_gpt":0.27308852827714064,"score_spread":0.26217701987616837,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1966038182","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9962308,0.00009003129,0.0017219298,0.00013625964,0.0000065669246,0.000056058016,0.00012133893,0.000021749871,0.0016152456],"genre_scores_gemma":[0.99773127,0.00006444951,0.0017364849,0.000011289478,0.000002150547,0.00004088542,0.00009617733,0.0000028293575,0.00031464107],"study_design_codex":"observational","study_design_gemma":"qualitative","domain_scores_codex":[0.99845076,0.000901978,0.000075209326,0.00010244995,0.00033948047,0.00013008721],"domain_scores_gemma":[0.9866861,0.009883043,0.0011024374,0.00045086144,0.001394322,0.00048321407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026800972,0.00024156031,0.00034716274,0.0033550553,0.000871274,0.0010680998,0.0003231936,0.00032340348,0.0012837419],"category_scores_gemma":[0.015058386,0.000106799846,0.0003531523,0.0023461361,0.00062065624,0.0011818644,0.00096824666,0.00042102026,0.00013707922],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016897721,0.00074471993,0.4641504,0.00040608246,0.00012899707,0.00090417656,0.2077726,0.0015928721,0.0020782659,0.0013287734,0.0013565378,0.31784678],"study_design_scores_gemma":[0.000120179706,0.0022388606,0.72385055,0.00017372286,0.00020385007,0.0012361237,0.24615175,0.0062753246,0.004487065,0.0043328283,0.010752133,0.00017761138],"about_ca_topic_score_codex":0.0058044745,"about_ca_topic_score_gemma":0.00590682,"teacher_disagreement_score":0.0058044745,"about_ca_system_score_codex":0.0012343475,"about_ca_system_score_gemma":0.0011318281,"threshold_uncertainty_score":0.014173865},"labels":[],"label_agreement":null},{"id":"W1966212893","doi":"10.1002/pmic.201090034","title":"Implementing Data Standards: A report on the HUPOPSI Workshop September 2009, Toronto, Canada","year":2010,"lang":"en","type":"article","venue":"PROTEOMICS","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Data science; Computer science","score_opus":0.02461406732161878,"score_gpt":0.3128680090074131,"score_spread":0.28825394168579427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1966212893","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062004272,0.03891896,0.028266726,0.62414783,0.0239612,0.0055696247,0.01987884,0.0024955184,0.194757],"genre_scores_gemma":[0.15902585,0.053088896,0.057933003,0.037622686,0.0025253934,0.002071505,0.03912394,0.002854005,0.64575464],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9751175,0.0030196665,0.0011393467,0.0018205285,0.013880541,0.0050223903],"domain_scores_gemma":[0.9496684,0.0037780812,0.0008294311,0.0016264829,0.03228611,0.011811454],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.045067217,0.001686534,0.0009013428,0.002269948,0.011910515,0.013571256,0.0048818137,0.004696738,0.024185829],"category_scores_gemma":[0.023702156,0.0014513701,0.0010629279,0.0045293584,0.0037666142,0.0046915384,0.0070571466,0.006842812,0.004982428],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011296449,0.00019424978,0.0032282749,0.00041810697,0.000021929636,0.00048944185,0.008191378,0.0008687628,0.0020648094,0.009041258,0.8965268,0.07884196],"study_design_scores_gemma":[0.000021810145,0.00003318072,0.008059688,0.00054221856,0.000011433895,0.000078886806,0.008969806,0.0004062234,0.0009646863,0.0006195537,0.98022527,0.00006733847],"about_ca_topic_score_codex":0.9395854,"about_ca_topic_score_gemma":0.9508172,"teacher_disagreement_score":0.9549328,"about_ca_system_score_codex":0.0820271,"about_ca_system_score_gemma":0.23199148,"threshold_uncertainty_score":0.59515107},"labels":[],"label_agreement":null},{"id":"W1967653863","doi":"10.1016/j.jbi.2012.09.006","title":"A survey of SNOMED CT implementations","year":2012,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":152,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"University of Victoria","keywords":"SNOMED CT; Terminology; Systematized Nomenclature of Medicine; Implementation; Computer science; Medicine; Medical physics; Software engineering","score_opus":0.04487624873121079,"score_gpt":0.34734098875464736,"score_spread":0.3024647400234366,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1967653863","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024384039,0.042370945,0.8232711,0.004736366,0.0005303479,0.00091291533,0.022196881,0.048579868,0.033017553],"genre_scores_gemma":[0.050341617,0.03472273,0.83970547,0.0022755996,0.00014458831,0.0005490181,0.052663345,0.0107825035,0.008815189],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9934316,0.0013751301,0.0013714067,0.00088536995,0.0026065959,0.00032997518],"domain_scores_gemma":[0.9826333,0.0076762694,0.00084944203,0.004560209,0.0039045834,0.00037615764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011160326,0.002033637,0.0015751648,0.014546032,0.0014507545,0.006172726,0.006750308,0.00209483,0.010384208],"category_scores_gemma":[0.028847706,0.0023937821,0.002775052,0.020570332,0.0012197421,0.008705979,0.00383525,0.0034674841,0.0042203125],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000756539,0.00040406163,0.0054403604,0.004368945,0.00053088815,0.0003127302,0.0007260952,0.0077255857,0.006680682,0.109649874,0.062120114,0.8012841],"study_design_scores_gemma":[0.00023738244,0.00041604522,0.005288769,0.006501039,0.00082409376,0.0025849054,0.000761676,0.059123907,0.025355877,0.0980768,0.80050266,0.00032685505],"about_ca_topic_score_codex":0.010029315,"about_ca_topic_score_gemma":0.014433316,"teacher_disagreement_score":0.014546032,"about_ca_system_score_codex":0.0024666602,"about_ca_system_score_gemma":0.006025569,"threshold_uncertainty_score":0.05902219},"labels":[],"label_agreement":null},{"id":"W1969581154","doi":"10.1145/2484028.2484167","title":"Exploiting semantics for improving clinical information retrieval","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Information retrieval; Query expansion; Ontology; Ranking (information retrieval); Web query classification; Query language; Semantics (computer science); Query optimization; Context (archaeology); Sargable; RDF query language; Web search query; Search engine","score_opus":0.02924178395879767,"score_gpt":0.3137144416713614,"score_spread":0.28447265771256375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969581154","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06582524,0.009502135,0.910869,0.0032959045,0.00020593764,0.00035314547,0.0011639389,0.0025589191,0.006225778],"genre_scores_gemma":[0.5189617,0.0050196364,0.47080353,0.00087361946,0.00036382137,0.0001884042,0.0024946665,0.00022052051,0.0010740241],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951709,0.0023964054,0.00065091986,0.00044624388,0.00118419,0.00015135114],"domain_scores_gemma":[0.99343115,0.0033507745,0.00084665755,0.0010260275,0.0012169954,0.00012850166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004761541,0.0008437435,0.0010102888,0.0064304816,0.0005916548,0.0021636898,0.00081438955,0.00095468486,0.0011640111],"category_scores_gemma":[0.017987346,0.00037634754,0.0011716573,0.005248166,0.0009218187,0.007429944,0.00226885,0.0008904209,0.00067101174],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005629375,0.00046382798,0.010601349,0.00163843,0.0002821739,0.0004865007,0.0010200007,0.04111703,0.04054672,0.07262294,0.011835832,0.81882226],"study_design_scores_gemma":[0.00031277168,0.00082248397,0.006886702,0.00045888103,0.0007326959,0.0023485075,0.0017435998,0.46717414,0.039712388,0.3994269,0.08009105,0.0002898912],"about_ca_topic_score_codex":0.0021674559,"about_ca_topic_score_gemma":0.0023599088,"teacher_disagreement_score":0.0064304816,"about_ca_system_score_codex":0.0010956576,"about_ca_system_score_gemma":0.0020415247,"threshold_uncertainty_score":0.02518177},"labels":[],"label_agreement":null},{"id":"W1972136887","doi":"10.1002/meet.14504701421","title":"Expediting medical literature coding with query‐building","year":2010,"lang":"en","type":"article","venue":"Proceedings of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"NeuroDevNet; University of British Columbia","funders":"","keywords":"Computer science; Expediting; Information retrieval; Bigram; Set (abstract data type); Haystack; Natural language processing; Artificial intelligence","score_opus":0.00461834009426845,"score_gpt":0.25783956931122914,"score_spread":0.2532212292169607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1972136887","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021931726,0.0010267028,0.866558,0.0027966963,0.00038658408,0.012359046,0.035230763,0.052556586,0.007153925],"genre_scores_gemma":[0.04377357,0.0003672241,0.92874223,0.00034797896,0.000107247186,0.005479738,0.017921189,0.0017415334,0.0015193503],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96526515,0.013138575,0.010713794,0.004162832,0.005882879,0.0008367901],"domain_scores_gemma":[0.7939723,0.13095558,0.01234111,0.025633818,0.034300204,0.0027969778],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04805315,0.0020419436,0.0020786242,0.03170725,0.002620945,0.0057352255,0.0031078602,0.001249585,0.02116337],"category_scores_gemma":[0.17235889,0.0012718156,0.0020503073,0.019091796,0.001856253,0.0057798657,0.008359228,0.001907683,0.0076485877],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012107354,0.00024866272,0.008622167,0.009428413,0.0002587106,0.0007819529,0.01229747,0.0030735,0.045535985,0.025622493,0.10700599,0.78591394],"study_design_scores_gemma":[0.00094646553,0.00079411396,0.02670828,0.0035036479,0.0006684324,0.0023060574,0.018395185,0.13860789,0.12088932,0.10077258,0.58548075,0.00092723203],"about_ca_topic_score_codex":0.008218818,"about_ca_topic_score_gemma":0.008131459,"teacher_disagreement_score":0.04805315,"about_ca_system_score_codex":0.0030808495,"about_ca_system_score_gemma":0.011260901,"threshold_uncertainty_score":0.25413233},"labels":[],"label_agreement":null},{"id":"W1973391887","doi":"10.1007/s10278-013-9598-3","title":"A Novel Knowledge Representation Framework for the Statistical Validation of Quantitative Imaging Biomarkers","year":2013,"lang":"en","type":"review","venue":"Journal of Digital Imaging","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Research (Canada); Bamfield Marine Sciences Centre","funders":"National Institute of Standards and Technology","keywords":"Computer science; Data science; Pipeline (software); Profiling (computer programming); Ontology; Annotation; Informatics; Information retrieval; Data mining; Artificial intelligence","score_opus":0.09909668672444769,"score_gpt":0.43153141329479594,"score_spread":0.3324347265703482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1973391887","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027113897,0.05267301,0.93631446,0.0020539144,0.00023726498,0.00020543601,0.0017111588,0.0014723866,0.0026209007],"genre_scores_gemma":[0.073454306,0.06277208,0.8502646,0.0014524574,0.00048468975,0.0005526775,0.008429114,0.00019012255,0.0023998301],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99597305,0.0009186347,0.00060604105,0.00069549016,0.0016793752,0.00012740195],"domain_scores_gemma":[0.99344325,0.0035887908,0.0006791975,0.0007804867,0.0014107645,0.00009752209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071199555,0.0013128596,0.0022104785,0.007873741,0.0005978084,0.004427285,0.004031025,0.0019454216,0.000943648],"category_scores_gemma":[0.013037451,0.00046828875,0.0020755117,0.0074980445,0.0014737836,0.0043585836,0.002556423,0.002237328,0.0007692431],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005658434,0.0001112367,0.0015195218,0.0024310737,0.0004041759,0.00023733098,0.00019732687,0.009562922,0.0027762197,0.057238434,0.009774138,0.915691],"study_design_scores_gemma":[0.0000702786,0.00013517801,0.0071528465,0.005101307,0.0013765551,0.0026515569,0.0004357287,0.21014823,0.014594669,0.5012881,0.25675857,0.00028700437],"about_ca_topic_score_codex":0.006193794,"about_ca_topic_score_gemma":0.0044218763,"teacher_disagreement_score":0.007873741,"about_ca_system_score_codex":0.0018662065,"about_ca_system_score_gemma":0.0039163767,"threshold_uncertainty_score":0.03765434},"labels":[],"label_agreement":null},{"id":"W1974493848","doi":"10.1191/0962280202sm278ra","title":"Analysis of repeated events","year":2002,"lang":"en","type":"review","venue":"Statistical Methods in Medical Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":143,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research","keywords":"Computer science; Event (particle physics); Marginal model; Medicine; Machine learning; Regression analysis","score_opus":0.3234065665027732,"score_gpt":0.6439255386965067,"score_spread":0.3205189721937335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974493848","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013612753,0.07316925,0.8951995,0.0027457583,0.0010488627,0.00080751366,0.0041727563,0.000919044,0.00832465],"genre_scores_gemma":[0.2360952,0.08651036,0.6534034,0.0023164141,0.0029021984,0.0029208546,0.009251489,0.00046076206,0.006139251],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98127466,0.0093132965,0.0019668182,0.003163727,0.004041609,0.00023986663],"domain_scores_gemma":[0.8902195,0.08951974,0.009412881,0.005992922,0.004437337,0.0004177062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020647401,0.0015755022,0.0028347864,0.0060039517,0.0006297487,0.003316192,0.002371321,0.0013053027,0.0048396327],"category_scores_gemma":[0.07056534,0.00061288424,0.0031470798,0.006543427,0.0018374659,0.0041745356,0.0017900962,0.0019713102,0.001349505],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034436796,0.00012711211,0.03061307,0.010481212,0.0033622808,0.0017112027,0.001217795,0.0068995017,0.0024549607,0.045921642,0.011083175,0.8857836],"study_design_scores_gemma":[0.00018703929,0.0007916744,0.08193152,0.0069206264,0.0048646415,0.010046283,0.0020969082,0.039198034,0.008977655,0.5870699,0.257496,0.00041973556],"about_ca_topic_score_codex":0.001250864,"about_ca_topic_score_gemma":0.0012456616,"teacher_disagreement_score":0.020647401,"about_ca_system_score_codex":0.0010074098,"about_ca_system_score_gemma":0.002128423,"threshold_uncertainty_score":0.10919523},"labels":[],"label_agreement":null},{"id":"W1974896839","doi":"10.1093/bib/bbv011","title":"The role of ontologies in biological and biomedical research: a functional perspective","year":2015,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":279,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Association of Occupational Therapists; Saint Paul University","funders":"","keywords":"Computer science; Biomedicine; Open Biomedical Ontologies; Metadata; Ontology; Controlled vocabulary; Data science; Domain (mathematical analysis); Perspective (graphical); Meaning (existential); Identifier; Vocabulary; Biological data; IDEF5; Information retrieval; Knowledge management; Artificial intelligence; Domain knowledge; World Wide Web; Ontology-based data integration; Ontology alignment; Epistemology; Bioinformatics","score_opus":0.09409572770032755,"score_gpt":0.33895458729472977,"score_spread":0.24485885959440223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1974896839","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01047401,0.07021101,0.6417508,0.17163838,0.0022520996,0.00016274468,0.0005660798,0.0003902467,0.10255468],"genre_scores_gemma":[0.39754194,0.06868283,0.5008624,0.01378288,0.005460973,0.0006358714,0.000989999,0.00032198068,0.011721063],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9830618,0.011424348,0.0013501323,0.0010828363,0.0024560562,0.00062480674],"domain_scores_gemma":[0.965452,0.024863478,0.0021450792,0.00330342,0.0031099583,0.0011259693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025977314,0.0011434833,0.0013531059,0.009514142,0.0047915205,0.014888771,0.0034295896,0.0048386795,0.002354599],"category_scores_gemma":[0.020401634,0.00096612034,0.0014488044,0.011034314,0.036790058,0.036124118,0.005635319,0.0063855303,0.00074220786],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000046363816,0.000007681444,0.00017045092,0.00011769376,0.000012285147,0.00006322125,0.001295939,0.0002649269,0.00013342607,0.98857576,0.0011108957,0.008243191],"study_design_scores_gemma":[0.000006006949,0.000010167103,0.00030098818,0.00044755675,0.000023658733,0.00019661056,0.0015273821,0.00110511,0.0002204302,0.9010645,0.0950771,0.000020454043],"about_ca_topic_score_codex":0.006353466,"about_ca_topic_score_gemma":0.0035131406,"teacher_disagreement_score":0.025977314,"about_ca_system_score_codex":0.0053149317,"about_ca_system_score_gemma":0.0056901006,"threshold_uncertainty_score":0.1373828},"labels":[],"label_agreement":null},{"id":"W1976357995","doi":"10.1145/1390334.1390502","title":"A reranking model for genomics aspect search","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Genomics; Domain (mathematical analysis); Artificial intelligence; Genome; Biology; Mathematics; Gene","score_opus":0.06892813578759659,"score_gpt":0.3042661842231046,"score_spread":0.23533804843550798,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1976357995","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030036308,0.00078616326,0.9630801,0.00041746485,0.00008611999,0.00015617875,0.00033090296,0.0032686312,0.0018380379],"genre_scores_gemma":[0.41965953,0.0006192246,0.56903,0.00032831242,0.0002694941,0.00027018943,0.0019129523,0.00039103834,0.0075193294],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986808,0.00042878248,0.00010407934,0.00028019003,0.00040629852,0.00009991083],"domain_scores_gemma":[0.99716634,0.0014334858,0.00021809895,0.0003318851,0.00075017044,0.00009991956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018419893,0.0009568513,0.0013420979,0.0017290282,0.0005783874,0.0011230984,0.0017774979,0.0012687196,0.0019078718],"category_scores_gemma":[0.005966881,0.0003806616,0.0007539148,0.001826128,0.00045404557,0.002841076,0.00062274875,0.0014041989,0.0012429245],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005735624,0.0003788977,0.0044313665,0.00026362124,0.00016108497,0.00033457222,0.00038504766,0.32391468,0.015320443,0.018646771,0.012761877,0.62282807],"study_design_scores_gemma":[0.000026844613,0.000061479164,0.00024000753,0.0000053818203,0.00001949814,0.000069962705,0.0000141511355,0.98858404,0.0022043902,0.006999064,0.0017623539,0.000012752047],"about_ca_topic_score_codex":0.008434837,"about_ca_topic_score_gemma":0.0133166,"teacher_disagreement_score":0.008434837,"about_ca_system_score_codex":0.0008418851,"about_ca_system_score_gemma":0.0011862189,"threshold_uncertainty_score":0.016771495},"labels":[],"label_agreement":null},{"id":"W197766533","doi":"","title":"Does GEM-encoding clinical practice guidelines improve the quality of knowledge bases? A study with the rule-based formalism.","year":2003,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Guideline; Formalism (music); Clinical Practice; Knowledge base; Computer science; Medicine; Data mining; Artificial intelligence; Family medicine","score_opus":0.09911141803199154,"score_gpt":0.4025723986414182,"score_spread":0.30346098060942667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W197766533","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.085680515,0.0015416044,0.9006154,0.003006823,0.000104063605,0.0005865928,0.0017383305,0.002907994,0.0038188205],"genre_scores_gemma":[0.37615603,0.00091029453,0.61866456,0.0004880073,0.00003553992,0.00018668227,0.00273076,0.00019365178,0.0006345003],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98235804,0.011129806,0.0021786964,0.0011513775,0.0028916944,0.0002903469],"domain_scores_gemma":[0.85055023,0.11657747,0.0071910885,0.015642446,0.009348497,0.00069026026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.036768097,0.00060110557,0.0011549164,0.0036527836,0.0006653329,0.004866542,0.0020404789,0.0013294979,0.0018908358],"category_scores_gemma":[0.16405544,0.0007429891,0.0014609944,0.004705513,0.0014294173,0.008479428,0.0024891745,0.0019973894,0.000375878],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012781175,0.0006206923,0.030313402,0.0022245236,0.00073578826,0.00061325746,0.0038979661,0.1456378,0.006942782,0.11856404,0.00483212,0.6843396],"study_design_scores_gemma":[0.000545239,0.0007537412,0.010379081,0.0018387063,0.0010928605,0.0010149257,0.0019393698,0.7540588,0.01886859,0.169311,0.040009752,0.00018793583],"about_ca_topic_score_codex":0.010576899,"about_ca_topic_score_gemma":0.010484727,"teacher_disagreement_score":0.036768097,"about_ca_system_score_codex":0.0026625039,"about_ca_system_score_gemma":0.0034748774,"threshold_uncertainty_score":0.19445062},"labels":[],"label_agreement":null},{"id":"W1978135684","doi":"10.1016/j.jcrc.2009.08.008","title":"Systematized Nomenclature of Medicine–Clinical Terms direction and its implications on critical care","year":2009,"lang":"en","type":"article","venue":"Journal of Critical Care","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":33,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Alberta Health Services","funders":"U.S. National Library of Medicine","keywords":"SNOMED CT; Terminology; Systematized Nomenclature of Medicine; Documentation; Medicine; Health informatics; Health care; Information system; Nursing; Computer science; Public health","score_opus":0.04167848993044548,"score_gpt":0.4193500061010354,"score_spread":0.3776715161705899,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1978135684","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027751746,0.0018483816,0.9251849,0.014577377,0.0022148397,0.00080713176,0.004981919,0.0018331975,0.020800529],"genre_scores_gemma":[0.14670447,0.0008102609,0.84130806,0.0016911114,0.0003268353,0.00039878365,0.0060360976,0.00031771557,0.0024066872],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9868774,0.006110314,0.003712708,0.0012335159,0.0017445585,0.00032140096],"domain_scores_gemma":[0.96757686,0.013978176,0.0032479395,0.004574028,0.0094384095,0.001184521],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012656128,0.00042881278,0.0006166759,0.005840072,0.0022971248,0.0066407453,0.001938861,0.0014520654,0.0036598048],"category_scores_gemma":[0.030089863,0.0006579587,0.0010259901,0.007979603,0.0036073306,0.008028151,0.0027327402,0.0033761184,0.0016426601],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012424511,0.00014213899,0.006497003,0.0005681495,0.0000753216,0.00040748718,0.0034760006,0.003442153,0.0054107006,0.84732586,0.019053452,0.11347759],"study_design_scores_gemma":[0.00016927223,0.00010917422,0.0066248593,0.0009396595,0.00030791198,0.0017191198,0.004150033,0.040496826,0.009781298,0.60315263,0.33239636,0.00015294907],"about_ca_topic_score_codex":0.008736064,"about_ca_topic_score_gemma":0.008846514,"teacher_disagreement_score":0.012656128,"about_ca_system_score_codex":0.004518664,"about_ca_system_score_gemma":0.0142127685,"threshold_uncertainty_score":0.0669328},"labels":[],"label_agreement":null},{"id":"W1980461251","doi":"10.2196/medinform.3023","title":"OWLing Clinical Data Repositories With the Ontology Web Language","year":2014,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; World Wide Web; Ontology; Unified Medical Language System; Information retrieval","score_opus":0.026548064786147137,"score_gpt":0.3563886985409111,"score_spread":0.32984063375476397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1980461251","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029027087,0.0003356192,0.976217,0.0010014967,0.00013021854,0.00047827762,0.0017272695,0.011310622,0.0058968626],"genre_scores_gemma":[0.033807725,0.0007505666,0.9517447,0.0006309468,0.00009598597,0.00047919856,0.007190156,0.0018510754,0.0034496887],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99322975,0.0013922103,0.0015404383,0.0007841859,0.0026852298,0.00036821255],"domain_scores_gemma":[0.99434304,0.0018480134,0.0007378129,0.0019402294,0.0007777585,0.00035328948],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007805538,0.00087863917,0.00076019525,0.005218695,0.0013132835,0.0065760673,0.0029585862,0.0017789545,0.0035214953],"category_scores_gemma":[0.013623604,0.0010777962,0.003049841,0.004313326,0.0019349286,0.00964108,0.0069288244,0.0027514137,0.0018801705],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002111534,0.0003387502,0.0031965978,0.0014970497,0.00027652414,0.0017331585,0.002197574,0.016608339,0.007830657,0.55760723,0.038654126,0.36984885],"study_design_scores_gemma":[0.00011474042,0.000096726864,0.0015925976,0.0011017165,0.00016287477,0.0019435487,0.00071773864,0.11992105,0.013000752,0.34937552,0.5117733,0.00019942023],"about_ca_topic_score_codex":0.011928053,"about_ca_topic_score_gemma":0.009790336,"teacher_disagreement_score":0.011928053,"about_ca_system_score_codex":0.0017574063,"about_ca_system_score_gemma":0.005691938,"threshold_uncertainty_score":0.04128015},"labels":[],"label_agreement":null},{"id":"W1982088273","doi":"10.1016/j.ymeth.2014.10.027","title":"Text as data: Using text-based features for proteins representation and for computational prediction of their characteristics","year":2014,"lang":"en","type":"review","venue":"Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Sinai Hospital; University of Toronto; Queen's University","funders":"U.S. National Library of Medicine; Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Representation (politics); Function (biology); Computational biology; Protein sequencing; Value (mathematics); Protein structure database; Process (computing); Genome; Sequence (biology); Protein function prediction; Protein function; Data mining; Information retrieval; Bioinformatics; Biology; Machine learning; Gene; Sequence database; Genetics; Peptide sequence","score_opus":0.22460454025643464,"score_gpt":0.4994441404292237,"score_spread":0.2748396001727891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982088273","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034272561,0.86522454,0.10339828,0.0041927216,0.0016267166,0.0003038285,0.006035774,0.0020182475,0.013772631],"genre_scores_gemma":[0.027292052,0.8315419,0.11675186,0.0029774806,0.0016701533,0.0005612238,0.011114919,0.00038222366,0.007708097],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991535,0.00016867652,0.00009384126,0.0001769847,0.00037396638,0.000033042223],"domain_scores_gemma":[0.9976604,0.0016697844,0.00017661559,0.00012603574,0.00030768337,0.00005953967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014333233,0.001328608,0.0017083236,0.0057779863,0.00022433636,0.0025153977,0.0023013514,0.0015560256,0.0035365853],"category_scores_gemma":[0.004330429,0.000363276,0.0012017041,0.0067376746,0.0008860197,0.004477598,0.00096550793,0.0016929489,0.0039416836],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000051096187,0.000063475716,0.0004625483,0.0073852967,0.00013732296,0.00012710317,0.00011062752,0.0010271678,0.0029681467,0.0081913695,0.02753678,0.951939],"study_design_scores_gemma":[0.000051452673,0.000110943874,0.0040594707,0.0059846938,0.00032909363,0.0015634102,0.00029202958,0.007016511,0.006840672,0.037866667,0.9357179,0.00016720216],"about_ca_topic_score_codex":0.0013620108,"about_ca_topic_score_gemma":0.0011662869,"teacher_disagreement_score":0.0057779863,"about_ca_system_score_codex":0.0007632763,"about_ca_system_score_gemma":0.00095828087,"threshold_uncertainty_score":0.011831045},"labels":[],"label_agreement":null},{"id":"W1982839620","doi":"10.1097/fpc.0b013e32832e2ced","title":"A systems biology network model for genetic association studies of nicotine addiction and treatment","year":2009,"lang":"en","type":"article","venue":"Pharmacogenetics and Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Addiction and Mental Health; University of Toronto","funders":"National Institute on Drug Abuse","keywords":"Ontology; Computer science; Genome-wide association study; Data science; Systems biology; Genetic association; Leverage (statistics); Bioinformatics; Artificial intelligence; Biology; Genetics","score_opus":0.04373480546693839,"score_gpt":0.3357536298030357,"score_spread":0.29201882433609727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1982839620","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010310625,0.0005926359,0.96660924,0.0042434516,0.00014860324,0.00020484473,0.003657641,0.00045590347,0.01377707],"genre_scores_gemma":[0.33270854,0.0023465778,0.64061916,0.0009694628,0.00020783431,0.0021147407,0.004951633,0.00013176363,0.015950324],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988574,0.0005442895,0.00007084627,0.0002780176,0.00019219537,0.000057334055],"domain_scores_gemma":[0.99742585,0.001834886,0.00022978966,0.00013398757,0.00026889125,0.00010657632],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020740593,0.0008494079,0.0005346345,0.0020673848,0.0009894557,0.0022375751,0.0018370505,0.0014797138,0.006596255],"category_scores_gemma":[0.0062322184,0.00040058073,0.002038592,0.002261142,0.0011510963,0.002398654,0.0014936371,0.0016870905,0.0007790984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007709702,0.00007284858,0.0031795746,0.00022784014,0.00016793609,0.00033976958,0.00045851464,0.42514107,0.0009912435,0.5410608,0.00510916,0.02317419],"study_design_scores_gemma":[0.00003735441,0.000021552498,0.00049571693,0.00004437901,0.00008294419,0.00010337852,0.000079270314,0.7340136,0.00013720659,0.2464418,0.018526271,0.00001644846],"about_ca_topic_score_codex":0.020181859,"about_ca_topic_score_gemma":0.024742667,"teacher_disagreement_score":0.020181859,"about_ca_system_score_codex":0.0033402517,"about_ca_system_score_gemma":0.003199977,"threshold_uncertainty_score":0.040128767},"labels":[],"label_agreement":null},{"id":"W1985592898","doi":"10.1186/1471-2105-12-s4-s6","title":"Deploying mutation impact text-mining software with the SADI Semantic Web Services framework","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"Natural Sciences and Engineering Research Council of Canada; New Brunswick Innovation Foundation; Canarie","keywords":"Computer science; World Wide Web; Context (archaeology); SPARQL; Web service; Semantic Web; Data science; Ontology; Reuse; Information retrieval; RDF; Biology","score_opus":0.021306098551409337,"score_gpt":0.2603210407245414,"score_spread":0.23901494217313204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1985592898","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040408824,0.00030759443,0.6510556,0.0021895359,0.00014066098,0.0012588432,0.011405466,0.28113803,0.01209551],"genre_scores_gemma":[0.16732313,0.00055474124,0.7846834,0.001088363,0.00007860704,0.00083847065,0.03011129,0.012154502,0.0031674274],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99720335,0.00035103157,0.00041175235,0.0004454527,0.0014281736,0.00016027926],"domain_scores_gemma":[0.9940724,0.002727233,0.0005772977,0.0012428019,0.0010289373,0.00035125983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061041457,0.00097065466,0.0007195738,0.003567818,0.0008654605,0.0025658596,0.0022272936,0.0009123677,0.0029429595],"category_scores_gemma":[0.010339154,0.00072759535,0.002080248,0.0019520099,0.0011331646,0.0031819905,0.0036461202,0.0018381648,0.0017218614],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021933212,0.0015176177,0.05219638,0.0041213804,0.0011722118,0.004494139,0.008292006,0.044608433,0.11503734,0.09012713,0.112501405,0.5637387],"study_design_scores_gemma":[0.00039567257,0.0002736964,0.012249374,0.00055365753,0.00041435452,0.002039154,0.001303513,0.4378112,0.1281177,0.075971045,0.34043726,0.00043343098],"about_ca_topic_score_codex":0.008307038,"about_ca_topic_score_gemma":0.008028181,"teacher_disagreement_score":0.008307038,"about_ca_system_score_codex":0.0023227911,"about_ca_system_score_gemma":0.003862672,"threshold_uncertainty_score":0.032282174},"labels":[],"label_agreement":null},{"id":"W1986657538","doi":"10.1093/nar/gkv383","title":"PolySearch2: a significantly improved text-mining system for discovering associations between human diseases, genes, drugs, metabolites, toxins and more","year":2015,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":148,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Genome Alberta; Canadian Institutes of Health Research; Alberta Innovates; Genome Canada","keywords":"DrugBank; UniProt; Information retrieval; Unified Medical Language System; Computer science; Computational biology; Ontology; Bioinformatics; Biology; Gene; Genetics; Drug","score_opus":0.07102885753606152,"score_gpt":0.373355959485395,"score_spread":0.3023271019493335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1986657538","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026013274,0.0024838585,0.23032731,0.0016058177,0.0003986018,0.0019389744,0.42180482,0.29185152,0.023575803],"genre_scores_gemma":[0.031523526,0.0014898314,0.42715544,0.00078797503,0.00016231702,0.001125593,0.5083677,0.007923426,0.021464102],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99837816,0.00025404844,0.00032249343,0.0004951094,0.00047043987,0.000079680336],"domain_scores_gemma":[0.99734384,0.0013966005,0.00025419894,0.00026801968,0.00054842234,0.00018892354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023348983,0.002955641,0.0013430461,0.009416708,0.0009971154,0.002025002,0.0016228975,0.0011458101,0.032975666],"category_scores_gemma":[0.0061632884,0.00073253026,0.0018864394,0.0054214355,0.00045700153,0.00409243,0.0025019727,0.0009357985,0.019030524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017287235,0.00051116344,0.013326233,0.0070173,0.0006696909,0.0024631878,0.0014214682,0.0032670796,0.035051636,0.010255983,0.46747443,0.45681307],"study_design_scores_gemma":[0.0006509479,0.00043290333,0.020930052,0.0007676319,0.0006115148,0.0033893776,0.0008672359,0.086595416,0.04748129,0.017307723,0.82064855,0.0003173638],"about_ca_topic_score_codex":0.0060603847,"about_ca_topic_score_gemma":0.0132111795,"teacher_disagreement_score":0.032975666,"about_ca_system_score_codex":0.00093579583,"about_ca_system_score_gemma":0.0028817921,"threshold_uncertainty_score":0.11031461},"labels":[],"label_agreement":null},{"id":"W1988270390","doi":"10.1503/cmaj.1050017","title":"Case summaries: another method","year":2005,"lang":"en","type":"article","venue":"Canadian Medical Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Information retrieval; Space (punctuation); Data science; World Wide Web; Data mining","score_opus":0.011763532402602252,"score_gpt":0.2775144238025777,"score_spread":0.26575089139997543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988270390","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008947833,0.0029842071,0.9216545,0.009665245,0.003371617,0.007503645,0.00821319,0.0036861342,0.033973645],"genre_scores_gemma":[0.0856515,0.0015413227,0.8815579,0.0017401665,0.0011192139,0.0039322586,0.003977752,0.0006143229,0.019865623],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9698224,0.013486907,0.0054351436,0.004650024,0.0059639597,0.0006415445],"domain_scores_gemma":[0.93847245,0.037857946,0.0037806896,0.008861841,0.009663765,0.0013633916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021307927,0.0017834004,0.001351451,0.020444414,0.0024497525,0.007395572,0.0033145233,0.0027474693,0.039758436],"category_scores_gemma":[0.09001064,0.00075702596,0.002931206,0.009687648,0.002533794,0.0073549077,0.00444102,0.0028141032,0.008426071],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078448356,0.00029599847,0.011189385,0.0029560286,0.00062571326,0.0013886399,0.0032933697,0.0013263049,0.0035262613,0.13851224,0.11351401,0.7225876],"study_design_scores_gemma":[0.0007760186,0.0003984578,0.008011902,0.0026027577,0.0012861369,0.013384155,0.0058388696,0.029374903,0.010040694,0.32324132,0.60454446,0.0005002945],"about_ca_topic_score_codex":0.0028856324,"about_ca_topic_score_gemma":0.0030342715,"teacher_disagreement_score":0.039758436,"about_ca_system_score_codex":0.0021875936,"about_ca_system_score_gemma":0.004309146,"threshold_uncertainty_score":0.13300526},"labels":[],"label_agreement":null},{"id":"W1988366874","doi":"10.1111/j.1467-8640.2011.00401.x","title":"EFFECTIVE BIO-EVENT EXTRACTION USING TRIGGER WORDS AND SYNTACTIC DEPENDENCIES","year":2011,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Negation; Natural language processing; Event (particle physics); Heuristics; Artificial intelligence; Syntax; Parsing; Task (project management); Biomedical text mining; Annotation; Scope (computer science); Machine learning; Programming language; Text mining","score_opus":0.05472830275994485,"score_gpt":0.33657597542470796,"score_spread":0.2818476726647631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988366874","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07804275,0.0013109454,0.8782843,0.001668914,0.00023963724,0.00085060846,0.009735822,0.020989692,0.008877316],"genre_scores_gemma":[0.23463675,0.0009853978,0.7396924,0.00044592988,0.00016523451,0.00044141186,0.019288056,0.0011921191,0.0031527232],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975101,0.0004505029,0.0004946021,0.00069051486,0.00073120004,0.00012315136],"domain_scores_gemma":[0.98998195,0.006595883,0.0011650722,0.0007027249,0.0013800337,0.00017429753],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027240887,0.0015452823,0.0012698342,0.005684592,0.0011879604,0.0024321768,0.0014547195,0.0014910365,0.005068121],"category_scores_gemma":[0.011287622,0.00067235285,0.0015263719,0.0030079973,0.00076187775,0.0052468763,0.0021684852,0.0016251148,0.003235477],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008958472,0.0004381663,0.017067757,0.002782919,0.0002348051,0.0030148898,0.0022807945,0.00895757,0.12743741,0.03917006,0.03785856,0.75986123],"study_design_scores_gemma":[0.00024538906,0.00028920788,0.019474972,0.00070856855,0.000667293,0.0048775543,0.002119848,0.41167212,0.29295537,0.13918442,0.12750275,0.00030251578],"about_ca_topic_score_codex":0.0013075311,"about_ca_topic_score_gemma":0.002366519,"teacher_disagreement_score":0.005684592,"about_ca_system_score_codex":0.0010889895,"about_ca_system_score_gemma":0.0024709206,"threshold_uncertainty_score":0.016954541},"labels":[],"label_agreement":null},{"id":"W1988833127","doi":"10.1016/j.vetpar.2006.02.018","title":"The impact on database searching arising from inconsistency in the nomenclature of parasitic diseases","year":2006,"lang":"en","type":"article","venue":"Veterinary Parasitology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute for Oil Sands Innovation, University of Alberta; Szent István Egyetem","keywords":"Nomenclature; Terminology; Disease; Coining (mint); Biology; Taxon; Bulgarian; Taxonomy (biology); Information retrieval; Computer science; Zoology; Medicine; Geography; Linguistics; Pathology; Ecology","score_opus":0.020767180004728953,"score_gpt":0.3510585639061173,"score_spread":0.33029138390138835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1988833127","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8675822,0.0077601895,0.07501553,0.02724485,0.00066492514,0.00031309717,0.0045123436,0.002658415,0.014248535],"genre_scores_gemma":[0.9133216,0.002564338,0.07317732,0.003136574,0.00046657346,0.00010067185,0.0044613606,0.0012395487,0.0015320198],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9116229,0.044893432,0.009035715,0.00539883,0.027454767,0.001594417],"domain_scores_gemma":[0.27729246,0.6617744,0.020149408,0.020532,0.018622654,0.0016290868],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.059460506,0.00079135085,0.0014220732,0.00657827,0.0024137625,0.008218344,0.0031215006,0.0031686018,0.0025361462],"category_scores_gemma":[0.43198928,0.0010366644,0.0015136409,0.0129294945,0.0028691948,0.010635006,0.004068702,0.0035295698,0.0008602492],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0055956934,0.0010356536,0.41349053,0.0031713797,0.0011171461,0.0056600985,0.006599796,0.03128789,0.015350503,0.023922078,0.021032605,0.4717366],"study_design_scores_gemma":[0.0009066809,0.0018954022,0.267417,0.0033929904,0.0057438784,0.031826388,0.017428061,0.38587174,0.058157768,0.16969275,0.056957368,0.00071001396],"about_ca_topic_score_codex":0.006479337,"about_ca_topic_score_gemma":0.005972945,"teacher_disagreement_score":0.9405395,"about_ca_system_score_codex":0.0024882583,"about_ca_system_score_gemma":0.0044035525,"threshold_uncertainty_score":0.314461},"labels":[],"label_agreement":null},{"id":"W1989217037","doi":"10.1016/j.jbi.2008.10.002","title":"Automatic summarization of MEDLINE citations for evidence-based medical treatment: A topic-oriented evaluation","year":2008,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":77,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Automatic summarization; Computer science; Information retrieval; MEDLINE; Multi-document summarization; Baseline (sea); Point (geometry); Medical physics; Medicine","score_opus":0.07552773137439109,"score_gpt":0.3564017328932263,"score_spread":0.28087400151883524,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1989217037","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.65847176,0.0918505,0.13617055,0.006416256,0.0017465284,0.009863267,0.072562106,0.011102416,0.011816628],"genre_scores_gemma":[0.5140364,0.016201217,0.3585709,0.0005708906,0.0014727205,0.0022831596,0.10286665,0.000514991,0.0034830598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99332505,0.0025616782,0.0018826765,0.00054636155,0.0015317214,0.00015256196],"domain_scores_gemma":[0.93210524,0.046300236,0.004678229,0.0012844331,0.014537672,0.0010941336],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012552086,0.0016539469,0.002494291,0.0266569,0.0012931223,0.0035043508,0.0014572553,0.0017291799,0.004905843],"category_scores_gemma":[0.054935273,0.0004107258,0.0024941794,0.012315979,0.00034665424,0.0026903881,0.0013827757,0.0007099207,0.0013261215],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009756435,0.001595003,0.034587387,0.021128817,0.005579092,0.00073234126,0.0015145831,0.009038252,0.02581791,0.0012475463,0.030964246,0.85803837],"study_design_scores_gemma":[0.011641148,0.017509172,0.2935063,0.007652251,0.08670725,0.006030401,0.008028497,0.30975252,0.10579512,0.015753895,0.13659261,0.0010308912],"about_ca_topic_score_codex":0.0032831735,"about_ca_topic_score_gemma":0.0065792385,"teacher_disagreement_score":0.9874479,"about_ca_system_score_codex":0.001014962,"about_ca_system_score_gemma":0.004606484,"threshold_uncertainty_score":0.06638253},"labels":[],"label_agreement":null},{"id":"W1991769959","doi":"10.1016/j.procs.2012.06.151","title":"HAIKU: A Semantic Framework for Surveillance of Healthcare-Associated Infections","year":2012,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; University of New Brunswick; McGill University","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Canarie","keywords":"Computer science; Health care; Ontology; Scope (computer science); Vocabulary; Data science; Knowledge management","score_opus":0.021712531590591867,"score_gpt":0.31197461196890863,"score_spread":0.29026208037831674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1991769959","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007966413,0.0022263352,0.9519567,0.0020772554,0.00028015443,0.0010210389,0.0073074084,0.014378851,0.012785829],"genre_scores_gemma":[0.08104167,0.0027533337,0.8969085,0.0009433406,0.0001459432,0.00083877443,0.013451207,0.0006808541,0.0032363522],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963737,0.0009521513,0.00072330376,0.0006103432,0.0010315907,0.00030883716],"domain_scores_gemma":[0.9975253,0.0009338964,0.000322826,0.00043267623,0.00050429587,0.00028104344],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004953821,0.0012907463,0.0013872881,0.00636294,0.0025851892,0.0060116453,0.0027128973,0.0026545704,0.0022249362],"category_scores_gemma":[0.0069588316,0.00082852546,0.0036796574,0.0043557226,0.0019477144,0.0076584634,0.005223626,0.0019015019,0.0011832973],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035441222,0.0003921338,0.010834268,0.0024489076,0.00070228614,0.0024224531,0.004712232,0.040376414,0.0063285255,0.6556678,0.048444673,0.22731592],"study_design_scores_gemma":[0.000092738264,0.00010225746,0.0049842247,0.001239015,0.0005599403,0.0018158925,0.0018116953,0.17285089,0.0061607156,0.30028352,0.50985414,0.00024495364],"about_ca_topic_score_codex":0.03041889,"about_ca_topic_score_gemma":0.033672173,"teacher_disagreement_score":0.03041889,"about_ca_system_score_codex":0.0026089028,"about_ca_system_score_gemma":0.0078045814,"threshold_uncertainty_score":0.060483694},"labels":[],"label_agreement":null},{"id":"W1992468151","doi":"10.1087/09531510125100269","title":"Bioline Publications: how its evolution has mirrored the growth of the internet","year":2001,"lang":"en","type":"article","venue":"Learned Publishing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"The Internet; Psychology; Computer science; World Wide Web","score_opus":0.05499437159891251,"score_gpt":0.2695080177761493,"score_spread":0.21451364617723678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1992468151","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072599195,0.07935494,0.033272047,0.35708517,0.015691532,0.00020582152,0.00563361,0.0049720965,0.4311856],"genre_scores_gemma":[0.46708456,0.113299504,0.0825625,0.028521469,0.02331746,0.00025973626,0.008777489,0.00792116,0.2682561],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99074274,0.0024533947,0.00076978817,0.0010516003,0.0044717696,0.0005107336],"domain_scores_gemma":[0.84812075,0.070479356,0.015829792,0.012456126,0.037580635,0.015533302],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.016632838,0.0003876385,0.0005252868,0.0152845,0.003068985,0.022083262,0.0016561806,0.0020534846,0.02104996],"category_scores_gemma":[0.06492976,0.00035790345,0.00037868554,0.024101768,0.005969011,0.021062579,0.0041021495,0.0035429916,0.008235545],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014918634,0.0001321773,0.01048765,0.0011410295,0.00003600208,0.0007240126,0.008093696,0.00049420196,0.0021366451,0.22219783,0.16035277,0.5940548],"study_design_scores_gemma":[0.000009017976,0.000025245945,0.0042983275,0.0004888946,0.000009167032,0.0004587395,0.001355445,0.0002789576,0.00062438205,0.010543872,0.98187524,0.000032587035],"about_ca_topic_score_codex":0.0036810562,"about_ca_topic_score_gemma":0.004430754,"teacher_disagreement_score":0.9779167,"about_ca_system_score_codex":0.004295205,"about_ca_system_score_gemma":0.005242175,"threshold_uncertainty_score":0.08796394},"labels":[],"label_agreement":null},{"id":"W1994715541","doi":"10.2196/medinform.3172","title":"Next Generation Phenotyping Using the Unified Medical Language System","year":2014,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; National Institutes of Health; National Center for Research Resources; Medical College of Wisconsin","keywords":"Computer science; Unified Medical Language System; Natural language processing","score_opus":0.03375254166050839,"score_gpt":0.30767290919886364,"score_spread":0.27392036753835525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1994715541","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017619917,0.0009360491,0.9321933,0.0016327447,0.00019558016,0.0008748825,0.01995938,0.019228395,0.0073597576],"genre_scores_gemma":[0.07740687,0.00049260125,0.8969277,0.0004439726,0.00007367814,0.0008614155,0.021720285,0.0005926599,0.001480803],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9946366,0.0021763374,0.0007355896,0.0011740808,0.0011009463,0.00017643985],"domain_scores_gemma":[0.99175256,0.0032803887,0.0012363592,0.0016907305,0.001793414,0.00024659338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007818481,0.00068942853,0.0006440983,0.0067673116,0.00084745756,0.0037910596,0.0013291383,0.0007978232,0.0032079632],"category_scores_gemma":[0.017765855,0.00038394477,0.0016245387,0.0043011643,0.0006780625,0.0028165958,0.003509325,0.0012469342,0.0018087237],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006672842,0.00027799493,0.033720236,0.001589768,0.0005241801,0.0021586202,0.0059149233,0.033521872,0.013518852,0.1796048,0.07093603,0.65756553],"study_design_scores_gemma":[0.00012597837,0.000287732,0.02155597,0.0012411312,0.00045241194,0.0016816256,0.0016034917,0.32032573,0.023630373,0.21743478,0.411368,0.00029276244],"about_ca_topic_score_codex":0.009884511,"about_ca_topic_score_gemma":0.008967227,"teacher_disagreement_score":0.009884511,"about_ca_system_score_codex":0.00213596,"about_ca_system_score_gemma":0.004956482,"threshold_uncertainty_score":0.041348517},"labels":[],"label_agreement":null},{"id":"W1996477238","doi":"10.1109/icmla.2012.82","title":"Integrating Machine Learning Into a Medical Decision Support System to Address the Problem of Missing Patient Data","year":2012,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Interoperability; Machine learning; Missing data; Artificial intelligence; Ontology; Decision support system; Task (project management); Data mining; Data science; World Wide Web","score_opus":0.027550123300884584,"score_gpt":0.3170047980085048,"score_spread":0.2894546747076202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1996477238","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0137704,0.00043954994,0.9747404,0.0032751514,0.00013537031,0.00024821173,0.00040594468,0.005528661,0.0014563149],"genre_scores_gemma":[0.1581432,0.00026456174,0.8387475,0.00088847306,0.00014975549,0.00013859116,0.0007981945,0.00012723287,0.00074243656],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959837,0.0013921092,0.00051801366,0.00080117234,0.0011234754,0.00018160544],"domain_scores_gemma":[0.98358685,0.011445355,0.000972807,0.0015555116,0.0018973935,0.00054202543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010032714,0.0008271884,0.0015720945,0.0030120914,0.0012688338,0.0035311095,0.0024804946,0.0017483092,0.0027578557],"category_scores_gemma":[0.019939497,0.0005391758,0.0010589386,0.0024555402,0.0008293122,0.003798943,0.002605095,0.0024000714,0.0011982694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084910274,0.0009489894,0.020013321,0.00080516346,0.0007265303,0.0018934872,0.0013361498,0.096767254,0.011074013,0.025983812,0.015659463,0.8239427],"study_design_scores_gemma":[0.00019768346,0.00020470962,0.0023381817,0.00032288747,0.00027914147,0.0010503891,0.00039248395,0.8724072,0.01876297,0.07596215,0.02793291,0.00014926243],"about_ca_topic_score_codex":0.0037366608,"about_ca_topic_score_gemma":0.0046258694,"teacher_disagreement_score":0.010032714,"about_ca_system_score_codex":0.0011599615,"about_ca_system_score_gemma":0.003475546,"threshold_uncertainty_score":0.053058684},"labels":[],"label_agreement":null},{"id":"W1997058260","doi":"10.1371/journal.pone.0051070","title":"A Unified Anatomy Ontology of the Vertebrate Skeletal System","year":2012,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; Dalhousie University","funders":"Division of Emerging Frontiers; National Human Genome Research Institute; National Institutes of Health; National Evolutionary Synthesis Center; National Science Foundation","keywords":"Vertebrate; Ontology; Biology; Evolutionary biology; Skeletal structures; Computational biology; Anatomy; Gene; Genetics","score_opus":0.029443479188401796,"score_gpt":0.2387185095606598,"score_spread":0.209275030372258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997058260","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010388748,0.0015372158,0.9017285,0.0024313997,0.0005178731,0.0010735017,0.03723815,0.007104493,0.037980173],"genre_scores_gemma":[0.054174706,0.0029315935,0.86958396,0.0010481777,0.00025490962,0.0014305061,0.057460666,0.00082925,0.01228624],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99829644,0.00025382737,0.00043037365,0.00031569746,0.0005625014,0.0001412296],"domain_scores_gemma":[0.998064,0.00048163967,0.00022001915,0.00036774133,0.00069986656,0.0001668356],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020876573,0.0007292847,0.0007057919,0.006177707,0.0017821064,0.0029396706,0.0018279247,0.0012299564,0.0061860443],"category_scores_gemma":[0.0035105527,0.00068909384,0.0019859795,0.005440663,0.0013622034,0.005334216,0.0026143233,0.0021873845,0.0024440098],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009459409,0.00014513898,0.003970968,0.0015683265,0.000104860505,0.0009228046,0.00323667,0.007946804,0.016546352,0.6505247,0.083853856,0.23108493],"study_design_scores_gemma":[0.00003950198,0.000052650452,0.00542558,0.00062476663,0.00011296163,0.0012484036,0.0006951872,0.017955525,0.003209823,0.10109215,0.8694521,0.00009137519],"about_ca_topic_score_codex":0.020164149,"about_ca_topic_score_gemma":0.025233138,"teacher_disagreement_score":0.020164149,"about_ca_system_score_codex":0.0023372541,"about_ca_system_score_gemma":0.009312164,"threshold_uncertainty_score":0.0400936},"labels":[],"label_agreement":null},{"id":"W1997148915","doi":"10.1038/nmeth.2698","title":"Significance, P values and t-tests","year":2013,"lang":"en","type":"article","venue":"Nature Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":185,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"","keywords":"Biology; Computational biology","score_opus":0.013180922589495428,"score_gpt":0.3727160138569054,"score_spread":0.35953509126740996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997148915","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38982335,0.037619922,0.3439423,0.010209674,0.019534199,0.004074448,0.10771884,0.009890175,0.077187106],"genre_scores_gemma":[0.78636265,0.0039182752,0.15792021,0.0045261513,0.0024320008,0.0136066405,0.016400868,0.0019157337,0.01291741],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97785234,0.0077333604,0.0025988235,0.0045839255,0.0059243934,0.0013071414],"domain_scores_gemma":[0.8989725,0.090368755,0.0030284915,0.0043709856,0.0019842486,0.0012749544],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015335266,0.001750561,0.003912857,0.0062050456,0.00210155,0.002623056,0.002982758,0.0039643426,0.038431965],"category_scores_gemma":[0.099860564,0.00055149005,0.002482606,0.007418008,0.0036740517,0.0030472944,0.0020668202,0.008731696,0.0048621907],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.024378022,0.004893884,0.13380258,0.03631006,0.015713839,0.0086410465,0.0035772335,0.010779801,0.023692805,0.065588646,0.24147122,0.43115088],"study_design_scores_gemma":[0.004796019,0.01834211,0.33230335,0.0047250646,0.015159468,0.011890514,0.011706628,0.035334647,0.017968837,0.33504277,0.21164303,0.00108763],"about_ca_topic_score_codex":0.0007094723,"about_ca_topic_score_gemma":0.00087238254,"teacher_disagreement_score":0.98466474,"about_ca_system_score_codex":0.0008031716,"about_ca_system_score_gemma":0.0014475989,"threshold_uncertainty_score":0.12856776},"labels":[],"label_agreement":null},{"id":"W1997441005","doi":"10.1109/iccabs.2012.6182651","title":"Identifying cancer biomarkers by knowledge discovery from medical literature","year":2012,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Data science; Field (mathematics); Flexibility (engineering); Knowledge extraction; Domain (mathematical analysis); Information retrieval; Cancer; Information extraction; Range (aeronautics); Association (psychology); Data mining; Medicine; Engineering; Psychology","score_opus":0.016951537974874937,"score_gpt":0.32141990828158506,"score_spread":0.30446837030671015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997441005","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23748928,0.06856388,0.56606734,0.012051431,0.00090048835,0.0031905815,0.08252981,0.009609217,0.019597923],"genre_scores_gemma":[0.24054438,0.022026677,0.6828179,0.0007761694,0.00050036353,0.0009359471,0.05026889,0.0002065572,0.0019232007],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974528,0.00055036834,0.00062610063,0.00050254486,0.0007531498,0.00011503372],"domain_scores_gemma":[0.9898375,0.0069642365,0.0013018632,0.0005492485,0.001119229,0.00022793554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036681287,0.0013504332,0.0014890082,0.037528127,0.0010017509,0.0029936216,0.0014558807,0.0011948927,0.0033069332],"category_scores_gemma":[0.017479576,0.00054916897,0.0020784875,0.017132306,0.0006199666,0.00338365,0.0019554538,0.00094990735,0.0024325252],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000601563,0.00051104435,0.050651662,0.010496699,0.0013142283,0.0050673163,0.0009020584,0.011730022,0.024857981,0.009024125,0.020082386,0.864761],"study_design_scores_gemma":[0.0005806658,0.0012940074,0.10850289,0.008894237,0.008942196,0.015384593,0.0037263343,0.19715661,0.086720906,0.18308167,0.3850577,0.0006582447],"about_ca_topic_score_codex":0.0027375198,"about_ca_topic_score_gemma":0.005118841,"teacher_disagreement_score":0.037528127,"about_ca_system_score_codex":0.00085811777,"about_ca_system_score_gemma":0.0040660943,"threshold_uncertainty_score":0.019399166},"labels":[],"label_agreement":null},{"id":"W1997463836","doi":"10.1186/1471-2105-10-s8-i1","title":"Between proteins and phenotypes: annotation and interpretation of mutations","year":2009,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Annotation; Phenotype; Computational biology; Population; Relevance (law); Biology; Mutation; Ontology; Genetics; Computer science; Bioinformatics; Gene; Medicine","score_opus":0.014369674518912403,"score_gpt":0.2682694366990187,"score_spread":0.2538997621801063,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997463836","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042695973,0.011790206,0.84467906,0.0066312756,0.0010527252,0.0011593603,0.04712456,0.024261793,0.02060504],"genre_scores_gemma":[0.19317496,0.006554077,0.73818547,0.0020007172,0.00057406106,0.0006505864,0.05052581,0.0039358586,0.004398483],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99376243,0.001756341,0.001095355,0.0014421031,0.001736268,0.00020754531],"domain_scores_gemma":[0.99104017,0.0048241382,0.0013527481,0.0014182656,0.0011079551,0.0002567847],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005698211,0.0018411531,0.0014325831,0.00984331,0.00107258,0.005108659,0.0028619396,0.0020993722,0.0064947954],"category_scores_gemma":[0.018150348,0.0006577314,0.0016031356,0.008116954,0.0017119307,0.0041024885,0.00451118,0.0017413326,0.00290667],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012399292,0.00048539965,0.02914921,0.007051465,0.0008337199,0.013339774,0.008501736,0.017134216,0.044572458,0.07665678,0.08665514,0.7143802],"study_design_scores_gemma":[0.00014473987,0.00014825989,0.03213742,0.0034711573,0.00074820174,0.009967394,0.0032856234,0.07013856,0.024514295,0.34908304,0.5059645,0.00039683562],"about_ca_topic_score_codex":0.003507168,"about_ca_topic_score_gemma":0.0018092857,"teacher_disagreement_score":0.00984331,"about_ca_system_score_codex":0.0011293332,"about_ca_system_score_gemma":0.0021748284,"threshold_uncertainty_score":0.030135334},"labels":[],"label_agreement":null},{"id":"W1997654049","doi":"10.4137/cin.s1046","title":"Integration of Neuroimaging and Microarray Datasets through Mapping and Model-Theoretic semantic Decomposition of Unstructured phenotypes","year":2009,"lang":"en","type":"article","venue":"Cancer Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"U.S. National Library of Medicine; National Cancer Institute","keywords":"Neuroimaging; Computer science; Phenotype; Data science; Artificial intelligence; Computational biology; Data mining; Natural language processing; Neuroscience; Biology","score_opus":0.0156274234471108,"score_gpt":0.2953569683706621,"score_spread":0.2797295449235513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1997654049","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01712391,0.00016596254,0.9777329,0.00054113136,0.000025807416,0.00015628953,0.0013271213,0.0020770598,0.0008498769],"genre_scores_gemma":[0.13772935,0.00018510724,0.855612,0.00022446389,0.000028470646,0.00034495365,0.005448136,0.00020763754,0.00021994446],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916679,0.00402761,0.0009309822,0.0013936354,0.0017831398,0.00019676522],"domain_scores_gemma":[0.9816042,0.009906468,0.00173789,0.0052200705,0.0012853798,0.00024593438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012393446,0.0012597715,0.0017930663,0.010064628,0.00096984435,0.005018028,0.0020066777,0.0008672886,0.00064937474],"category_scores_gemma":[0.03440634,0.00067913055,0.002891813,0.011381195,0.0013855142,0.005842428,0.0050729923,0.0019132348,0.00028726461],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060546293,0.0008329214,0.03143525,0.0012205112,0.0013942298,0.0014383564,0.004951558,0.14875005,0.022545742,0.19791473,0.0088476045,0.58006364],"study_design_scores_gemma":[0.000052934673,0.000116855175,0.0061512734,0.00017260609,0.00030140998,0.00045947416,0.0013411875,0.47462156,0.00889953,0.49514443,0.012633013,0.0001057478],"about_ca_topic_score_codex":0.00473598,"about_ca_topic_score_gemma":0.0060467254,"teacher_disagreement_score":0.012393446,"about_ca_system_score_codex":0.0017884192,"about_ca_system_score_gemma":0.0026048447,"threshold_uncertainty_score":0.06554359},"labels":[],"label_agreement":null},{"id":"W1998028266","doi":"10.1371/journal.pone.0022006","title":"Interoperability between Biomedical Ontologies through Relation Expansion, Upper-Level Ontologies and Automatic Reasoning","year":2011,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council; European Commission; Natural Sciences and Engineering Research Council of Canada; U.S. Public Health Service; National Institutes of Health; European Bioinformatics Institute","keywords":"Computer science; Ontology; Open Biomedical Ontologies; Automated reasoning; Interoperability; IDEF5; Semantics (computer science); Consistency (knowledge bases); Class (philosophy); Ontology components; Inference; Knowledge representation and reasoning; Relation (database); Description logic; Upper ontology; Data science; Information retrieval; Semantic Web; Programming language; World Wide Web; Artificial intelligence; Suggested Upper Merged Ontology; Data mining","score_opus":0.1398158278491219,"score_gpt":0.281447782143253,"score_spread":0.1416319542941311,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998028266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010425426,0.00025850613,0.98166186,0.0012022221,0.000029366722,0.0001943327,0.0003077681,0.0020488612,0.0038717156],"genre_scores_gemma":[0.13065131,0.00038759416,0.8643837,0.00043362335,0.00005934977,0.00027864642,0.0017446545,0.0003973384,0.001663812],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9814135,0.007184204,0.0025958084,0.0024164142,0.0056779888,0.00071215216],"domain_scores_gemma":[0.9779409,0.012448653,0.001460018,0.006333757,0.0015608525,0.00025576432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016897399,0.00085351855,0.0012779246,0.00629574,0.0023106236,0.0068412726,0.0034170568,0.0016149537,0.002950703],"category_scores_gemma":[0.03541177,0.0012696021,0.004136737,0.0060319444,0.004495929,0.020451836,0.007670014,0.0038441103,0.0008699495],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013401319,0.00020929199,0.0025127248,0.0004300853,0.00018119675,0.000613812,0.0036300628,0.026972948,0.0070364363,0.7567163,0.0049923533,0.19657084],"study_design_scores_gemma":[0.000033104603,0.000017495566,0.00068970164,0.00014671564,0.00010454469,0.00023105239,0.00044016915,0.1150035,0.007022412,0.8505475,0.025714496,0.00004930944],"about_ca_topic_score_codex":0.00857736,"about_ca_topic_score_gemma":0.0074471524,"teacher_disagreement_score":0.016897399,"about_ca_system_score_codex":0.003139293,"about_ca_system_score_gemma":0.0039568087,"threshold_uncertainty_score":0.08936304},"labels":[],"label_agreement":null},{"id":"W1998043417","doi":"10.1016/j.jbi.2010.03.003","title":"Automatic indexing and retrieval of encounter-specific evidence for point-of-care support","year":2010,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"Canadian Institutes of Health Research","keywords":"Search engine indexing; Information retrieval; Computer science; Systematic review; Precision and recall; MEDLINE; Point (geometry); Recall; Evidence-based medicine; Clinical decision support system; Empirical evidence; Empirical research; Decision support system; Data mining; Psychology","score_opus":0.027068045658596048,"score_gpt":0.3227079864957413,"score_spread":0.29563994083714523,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1998043417","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46135598,0.025280919,0.36158553,0.008535244,0.0015519297,0.0035535786,0.10234173,0.012026079,0.023769015],"genre_scores_gemma":[0.59360564,0.006367283,0.33070666,0.0005493659,0.00058406574,0.00059972674,0.06457935,0.00030492197,0.0027029729],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975303,0.0004153669,0.00061639,0.00029389976,0.0009759275,0.00016825706],"domain_scores_gemma":[0.9851275,0.00753511,0.0016988739,0.0014926371,0.0035067792,0.00063915347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018948601,0.0007960363,0.0012402543,0.017314205,0.00083682925,0.0032113164,0.0017678814,0.0017736069,0.0058895582],"category_scores_gemma":[0.029348917,0.0003836899,0.0009775171,0.0086738495,0.0004132113,0.0028461274,0.0023489448,0.0008625684,0.0024579193],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001790476,0.0009186122,0.07114741,0.0052431296,0.00047752695,0.0035952313,0.001354884,0.0059372317,0.029369853,0.009819727,0.05333976,0.81700623],"study_design_scores_gemma":[0.001463401,0.0017105188,0.15268666,0.006685007,0.0052262587,0.02325047,0.0077390624,0.35094544,0.1225028,0.09899144,0.22814626,0.000652735],"about_ca_topic_score_codex":0.004023376,"about_ca_topic_score_gemma":0.0075895684,"teacher_disagreement_score":0.017314205,"about_ca_system_score_codex":0.00096599106,"about_ca_system_score_gemma":0.0040631313,"threshold_uncertainty_score":0.019702554},"labels":[],"label_agreement":null},{"id":"W1999259627","doi":"10.4018/jhisi.2010110303","title":"A Framework for Data and Mined Knowledge Interoperability in Clinical Decision Support Systems","year":2010,"lang":"en","type":"article","venue":"International Journal of Healthcare Information Systems and Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Interoperability; Clinical decision support system; Computer science; Decision support system; Guideline; Health care; Knowledge management; Quality (philosophy); Data science; Data mining; Medicine; World Wide Web","score_opus":0.05877489027119352,"score_gpt":0.4263281211931375,"score_spread":0.367553230921944,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999259627","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005580993,0.00019525817,0.99532974,0.0008439534,0.000034234235,0.00033548527,0.00020251347,0.00089176785,0.0016089294],"genre_scores_gemma":[0.011786803,0.00024329808,0.98583543,0.00019390775,0.000038441565,0.0005500703,0.00066529505,0.00007882764,0.000607828],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9805459,0.008471059,0.0041054036,0.0018547623,0.004307123,0.0007156906],"domain_scores_gemma":[0.9847883,0.008306267,0.0010192051,0.0033639849,0.0017177729,0.00080443587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.032134075,0.0015524855,0.0019552975,0.006347307,0.0029765067,0.014943758,0.0071932157,0.005354934,0.0037507568],"category_scores_gemma":[0.030469311,0.0019979475,0.0055419407,0.006916934,0.0059539895,0.012613589,0.01007686,0.0056164665,0.0017476518],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007431438,0.000143736,0.00069774006,0.00058830006,0.00018519312,0.00093595893,0.0020929687,0.02824617,0.0016490316,0.88173145,0.0049784645,0.07867665],"study_design_scores_gemma":[0.00010866853,0.00010925034,0.00030798212,0.0008293919,0.00016244089,0.00062644965,0.00073762826,0.20102586,0.0027970446,0.6713673,0.12180442,0.00012363792],"about_ca_topic_score_codex":0.01325191,"about_ca_topic_score_gemma":0.011779458,"teacher_disagreement_score":0.032134075,"about_ca_system_score_codex":0.003870764,"about_ca_system_score_gemma":0.008223577,"threshold_uncertainty_score":0.16994321},"labels":[],"label_agreement":null},{"id":"W1999302328","doi":"10.1186/2041-1480-5-11","title":"Benchmarking infrastructure for mutation text mining","year":2014,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"New Brunswick Innovation Foundation","keywords":"Computer science; Benchmarking; SPARQL; Information retrieval; Ontology; RDF; Data mining; Annotation; Benchmark (surveying); Data science; Semantic Web; Artificial intelligence","score_opus":0.008037719161641739,"score_gpt":0.26644384119059755,"score_spread":0.2584061220289558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999302328","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.051744524,0.0012822968,0.64531434,0.0024818226,0.00039896826,0.002905316,0.04321505,0.22748612,0.025171507],"genre_scores_gemma":[0.24740711,0.0009342615,0.45418712,0.00070531707,0.00016689775,0.0034382064,0.27363917,0.014952503,0.00456938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9754655,0.007200564,0.004746299,0.0037546614,0.0073491815,0.001483853],"domain_scores_gemma":[0.9505344,0.011940997,0.0035146566,0.016448462,0.0154760415,0.0020854357],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024773045,0.0017316582,0.0014699383,0.009684004,0.0023167492,0.0038805797,0.005354342,0.0017301769,0.007214346],"category_scores_gemma":[0.051926505,0.0007507184,0.0015837393,0.010173993,0.0013513531,0.009203847,0.0058466005,0.0020959019,0.005492264],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002763261,0.001988039,0.025236648,0.004649151,0.00054731924,0.0012946684,0.002468098,0.059395384,0.04418221,0.08209286,0.23332098,0.54206145],"study_design_scores_gemma":[0.00051800616,0.0010601605,0.022992766,0.0011485406,0.00035366882,0.001355556,0.0014497199,0.3993921,0.1278121,0.08909531,0.35429904,0.00052306213],"about_ca_topic_score_codex":0.0067862994,"about_ca_topic_score_gemma":0.0035723213,"teacher_disagreement_score":0.024773045,"about_ca_system_score_codex":0.003956407,"about_ca_system_score_gemma":0.0068540466,"threshold_uncertainty_score":0.13101393},"labels":[],"label_agreement":null},{"id":"W1999868121","doi":"10.1109/fgct.2012.6476567","title":"Building a diseases symptoms ontology for medical diagnosis: An integrative approach","year":2012,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Ontology; Disease; Computer science; Open Biomedical Ontologies; Medicine; Upper ontology; Data science; Ontology alignment; Information retrieval; Semantic Web; Pathology","score_opus":0.01662990926078102,"score_gpt":0.3240751234765247,"score_spread":0.30744521421574367,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1999868121","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0076712123,0.0009650294,0.975045,0.0031111299,0.00013090103,0.0006491523,0.0020501947,0.0016315659,0.008745824],"genre_scores_gemma":[0.05969719,0.0010533211,0.93137276,0.0006359734,0.00009559879,0.00032346172,0.005242322,0.0001691646,0.0014102471],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99718297,0.00062037626,0.0005890107,0.0005853395,0.00090059306,0.00012177676],"domain_scores_gemma":[0.9973822,0.00083621865,0.00029166276,0.00032438902,0.000934068,0.00023142569],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044180197,0.00091324525,0.0010554962,0.007999512,0.0014713949,0.003287525,0.0015017177,0.0011928292,0.002473917],"category_scores_gemma":[0.006149019,0.0005838138,0.0028121073,0.0043482245,0.0011265659,0.0055738064,0.0039761052,0.0017343714,0.0009263151],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024995953,0.00065265375,0.017126465,0.0025523265,0.0009774643,0.0020769679,0.00403284,0.016016155,0.023506254,0.28117526,0.026576791,0.62505686],"study_design_scores_gemma":[0.00013780018,0.00023531176,0.015979648,0.0023026855,0.0019597088,0.003979196,0.0047104787,0.1780843,0.02400699,0.3794527,0.38886982,0.00028137813],"about_ca_topic_score_codex":0.0045799143,"about_ca_topic_score_gemma":0.0059193154,"teacher_disagreement_score":0.007999512,"about_ca_system_score_codex":0.0019462968,"about_ca_system_score_gemma":0.0064257034,"threshold_uncertainty_score":0.023364961},"labels":[],"label_agreement":null},{"id":"W2000466942","doi":"10.1038/418589a","title":"Thinking laterally about genes","year":2002,"lang":"en","type":"article","venue":"Nature","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Gene; Computational biology; Biology; Evolutionary biology; Genetics","score_opus":0.009994032277814108,"score_gpt":0.2542910436762211,"score_spread":0.24429701139840695,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000466942","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022092106,0.0011152838,0.77371585,0.057884686,0.00094335206,0.00006205246,0.00085366867,0.00071187236,0.14262111],"genre_scores_gemma":[0.63659334,0.0027376008,0.28041705,0.013346532,0.0016379873,0.00027836795,0.0014703898,0.000526012,0.06299275],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99878067,0.00044502039,0.00006258242,0.00026077678,0.00036577205,0.00008510624],"domain_scores_gemma":[0.99694604,0.0017390461,0.00020850329,0.00051078614,0.00043182497,0.00016382367],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017748964,0.0006505753,0.00039324744,0.0015235709,0.0015315484,0.0038204256,0.001148767,0.0012675794,0.01203341],"category_scores_gemma":[0.008026393,0.00041619997,0.00096733245,0.0013611504,0.005202375,0.016886668,0.0027140006,0.0037020906,0.0023338844],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019895704,0.0000056799527,0.00029483455,0.000035326957,0.000012881209,0.00007160976,0.00043959534,0.0005410089,0.0006888663,0.9794095,0.003412442,0.015068294],"study_design_scores_gemma":[0.000003888455,0.0000031332254,0.000067782275,0.000015546064,0.00001328748,0.00004784225,0.00018882984,0.0017444149,0.00055627385,0.9836077,0.013746209,0.0000050543126],"about_ca_topic_score_codex":0.0022850991,"about_ca_topic_score_gemma":0.0023317125,"teacher_disagreement_score":0.01203341,"about_ca_system_score_codex":0.0014760569,"about_ca_system_score_gemma":0.0012378301,"threshold_uncertainty_score":0.040255725},"labels":[],"label_agreement":null},{"id":"W2000868265","doi":"10.1097/01.nnr.0000280638.01773.84","title":"Whither Knowledge Translation","year":2007,"lang":"en","type":"letter","venue":"Nursing Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institutes of Health Research","funders":"","keywords":"Knowledge translation; Value (mathematics); Translation (biology); Engineering ethics; Political science; Library science; Medicine; Knowledge management; Computer science; Engineering; Biology","score_opus":0.2451655229185323,"score_gpt":0.4814073626351537,"score_spread":0.2362418397166214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2000868265","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000052642776,0.0012005151,0.00027085017,0.9727472,0.02362879,0.0000055464766,0.0000127121975,0.000023134753,0.0020585242],"genre_scores_gemma":[0.0017573547,0.0017588277,0.00094712316,0.9600614,0.027382523,0.000029644516,0.000032710697,0.000050909654,0.007979555],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9579707,0.013370928,0.0045456733,0.0045583323,0.016610947,0.0029434112],"domain_scores_gemma":[0.8034027,0.13849618,0.0060527124,0.010974804,0.027360205,0.013713476],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.031445444,0.0010935429,0.0015667835,0.0017166178,0.009367175,0.013185339,0.0035227675,0.069487765,0.029490938],"category_scores_gemma":[0.16786776,0.0008735699,0.0021841512,0.0021108014,0.016223453,0.027549075,0.011605744,0.094358295,0.022833055],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022210706,0.00002698673,0.000069119764,0.00012432924,0.000009637292,0.00039260313,0.00039280573,0.00004441775,0.00011313118,0.02241232,0.95314187,0.023250619],"study_design_scores_gemma":[0.000026471727,0.000028552884,0.00011730419,0.00037077823,0.000008685876,0.0004231568,0.0007881484,0.00012723054,0.000108394735,0.03527746,0.96269196,0.000031815856],"about_ca_topic_score_codex":0.0045270543,"about_ca_topic_score_gemma":0.008357347,"teacher_disagreement_score":0.96855456,"about_ca_system_score_codex":0.009053582,"about_ca_system_score_gemma":0.016549429,"threshold_uncertainty_score":0.16630137},"labels":[],"label_agreement":null},{"id":"W2001612861","doi":"10.1053/j.ajkd.2009.11.026","title":"Optimal Search Filters for Renal Information in EMBASE","year":2010,"lang":"en","type":"article","venue":"American Journal of Kidney Diseases","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Western University","funders":"Canadian Institutes of Health Research; Kidney Foundation of Canada","keywords":"Medicine; MEDLINE; Intensive care medicine; Urology","score_opus":0.0063149688872334835,"score_gpt":0.27406178495461175,"score_spread":0.26774681606737827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2001612861","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37351394,0.022329422,0.50623775,0.004072093,0.00062458287,0.0013172554,0.06731335,0.012761792,0.011829792],"genre_scores_gemma":[0.4287691,0.004157729,0.4956751,0.00042872943,0.0002493239,0.0005126194,0.06272103,0.0006120944,0.0068742754],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9979825,0.00045997067,0.00039832626,0.00037309944,0.000539205,0.00024694504],"domain_scores_gemma":[0.9927602,0.005110174,0.0003976603,0.00035629998,0.0012199976,0.00015568148],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.002568464,0.00109347,0.0019583502,0.014529023,0.0011068997,0.0028967534,0.0009823584,0.0016817398,0.005279273],"category_scores_gemma":[0.015237477,0.000501889,0.0016091327,0.008512522,0.00034559058,0.002659846,0.0009677307,0.0007327076,0.0023419466],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033642743,0.0006465801,0.02131827,0.003579087,0.0007588262,0.0013507693,0.0007065652,0.026941983,0.030272014,0.0125913,0.0712324,0.82723796],"study_design_scores_gemma":[0.0008199112,0.0011477211,0.04810108,0.0018862427,0.003027998,0.003036468,0.0030661505,0.71571285,0.080499835,0.056923985,0.085519314,0.00025843288],"about_ca_topic_score_codex":0.020781029,"about_ca_topic_score_gemma":0.031637687,"teacher_disagreement_score":0.9974315,"about_ca_system_score_codex":0.0012949501,"about_ca_system_score_gemma":0.0053005847,"threshold_uncertainty_score":0.041320145},"labels":[],"label_agreement":null},{"id":"W2002372013","doi":"10.1038/labinvest.3700694","title":"Creation of a retrospective searchable neuropathologic database from print archives at Toronto's University Health Network","year":2007,"lang":"en","type":"article","venue":"Laboratory Investigation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University; University Health Network; University of Toronto","funders":"University of Toronto; World Health Organization","keywords":"Neuropathology; Medical diagnosis; Computer science; Prioritization; Categorization; Medicine; Library science; Medical physics; Information retrieval; Pathology; Artificial intelligence; Disease; Engineering","score_opus":0.0130433731780127,"score_gpt":0.25453824527001867,"score_spread":0.24149487209200596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002372013","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22385186,0.0010813581,0.027368555,0.0010473293,0.000083928644,0.0021651469,0.73184127,0.0063771745,0.0061834836],"genre_scores_gemma":[0.26409048,0.0012953781,0.07389479,0.00017800198,0.00005782749,0.0011450526,0.6563722,0.0005233491,0.0024429092],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988846,0.00014160568,0.00027613298,0.00030156336,0.00030781483,0.00008826043],"domain_scores_gemma":[0.99274796,0.0026614703,0.0009986174,0.0010786523,0.0019346006,0.00057871215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021435625,0.0006106569,0.0007129843,0.011783328,0.0010017873,0.002167531,0.0011862466,0.000539368,0.0064754365],"category_scores_gemma":[0.011994053,0.0004936485,0.0005343512,0.007865711,0.00035476798,0.0011236394,0.001244017,0.0005277847,0.0022038345],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023926971,0.0006983902,0.34859946,0.005732621,0.00054683606,0.011669668,0.0055885524,0.012526888,0.042021807,0.0077058463,0.20527431,0.35724297],"study_design_scores_gemma":[0.00058170786,0.0005049336,0.53823644,0.0013466475,0.0015508541,0.0070273406,0.0059040342,0.04920226,0.04056553,0.0030983482,0.35158536,0.0003966025],"about_ca_topic_score_codex":0.12457042,"about_ca_topic_score_gemma":0.13060045,"teacher_disagreement_score":0.12457042,"about_ca_system_score_codex":0.0030993405,"about_ca_system_score_gemma":0.008979126,"threshold_uncertainty_score":0.24769068},"labels":[],"label_agreement":null},{"id":"W2002463857","doi":"10.1016/s0140-6736(13)62228-x","title":"Reducing waste from incomplete or unusable reports of biomedical research","year":2014,"lang":"en","type":"article","venue":"The Lancet","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1359,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"Medical Research Council; National Health and Medical Research Council; Cancer Research UK","keywords":"Waste management; MEDLINE; Medicine; Business; Engineering; Political science; Law","score_opus":0.0983632052959929,"score_gpt":0.36700159057976456,"score_spread":0.26863838528377165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002463857","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22958735,0.025734521,0.6340629,0.04263147,0.0017614897,0.0012541418,0.022780802,0.015003565,0.027183795],"genre_scores_gemma":[0.41520864,0.0076989937,0.53949773,0.0037445254,0.0008752882,0.0004782388,0.022420311,0.0038480614,0.0062282113],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9121966,0.033960145,0.010900105,0.005191396,0.035889223,0.0018625611],"domain_scores_gemma":[0.43553948,0.35503677,0.05214125,0.09143501,0.06198978,0.003857692],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06080722,0.0015601076,0.002623822,0.027177174,0.0016315309,0.011608967,0.0046478137,0.0023051044,0.0034207935],"category_scores_gemma":[0.32106128,0.0015219338,0.0025179596,0.030807422,0.002462167,0.014576508,0.0061388253,0.0029258374,0.0025533657],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014621839,0.0005030583,0.057995778,0.0039602146,0.0010845629,0.0012444207,0.004198461,0.013365399,0.008112236,0.03363118,0.05167041,0.8227721],"study_design_scores_gemma":[0.00054133387,0.0006285095,0.05353762,0.009219078,0.0055698715,0.0053640353,0.010720309,0.1812841,0.09589122,0.34966636,0.28702852,0.00054904516],"about_ca_topic_score_codex":0.006810002,"about_ca_topic_score_gemma":0.0057873204,"teacher_disagreement_score":0.9391928,"about_ca_system_score_codex":0.0033308894,"about_ca_system_score_gemma":0.01017956,"threshold_uncertainty_score":0.32158315},"labels":[],"label_agreement":null},{"id":"W2002507929","doi":"10.1016/j.tree.2007.03.013","title":"Phenotype ontologies: the bridge between genomics and evolution","year":2007,"lang":"en","type":"article","venue":"Trends in Ecology & Evolution","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":137,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of British Columbia","funders":"National Human Genome Research Institute; Johns Hopkins University; National Institutes of Health; National Science Foundation","keywords":"Genomics; Evolutionary developmental biology; Parallels; Biology; Evolutionary biology; Phenotype; Comparative genomics; Ontology; Gene; Genome; Genetics","score_opus":0.023001214715074776,"score_gpt":0.2948531741314533,"score_spread":0.2718519594163785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2002507929","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077654244,0.0027905796,0.96596855,0.0084130205,0.00041515185,0.00012831182,0.00340819,0.003004505,0.008106344],"genre_scores_gemma":[0.14418228,0.0060986117,0.8248299,0.0037097838,0.00064395595,0.0004129954,0.013836089,0.0013544116,0.0049320944],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9967565,0.0011956689,0.00038616493,0.00069564703,0.0008440138,0.000121967896],"domain_scores_gemma":[0.98602563,0.00845291,0.0013119739,0.0024251956,0.0010474552,0.0007368419],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00576526,0.00088880834,0.0012114333,0.005512386,0.0016312025,0.006146033,0.0030226258,0.0021575955,0.0041360296],"category_scores_gemma":[0.019058034,0.00073889806,0.0017855444,0.005458863,0.0036362663,0.016885743,0.0046142894,0.0038641158,0.0014497843],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012511214,0.00015332307,0.0040713097,0.0010210937,0.00019352298,0.00042556686,0.0018001798,0.0040574432,0.0036843657,0.72866845,0.01965027,0.23614947],"study_design_scores_gemma":[0.00002381611,0.000021529404,0.0013320203,0.0004187446,0.00009463854,0.00032836205,0.00046269494,0.019615093,0.0016313943,0.8850313,0.09100343,0.00003700783],"about_ca_topic_score_codex":0.0028696381,"about_ca_topic_score_gemma":0.0029877126,"teacher_disagreement_score":0.006146033,"about_ca_system_score_codex":0.0017808498,"about_ca_system_score_gemma":0.0035022234,"threshold_uncertainty_score":0.030489981},"labels":[],"label_agreement":null},{"id":"W2003185534","doi":"10.3115/1572340.1572361","title":"Syntactic dependency based heuristics for biological event extraction","year":2009,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":123,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Negation; Heuristics; Parsing; Dependency (UML); Biomedical text mining; Task (project management); Syntax; Dependency grammar; Event (particle physics); Natural language processing; Artificial intelligence; Programming language; Text mining","score_opus":0.030486956664318395,"score_gpt":0.3342910224198971,"score_spread":0.3038040657555787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003185534","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005759793,0.00044453674,0.9785156,0.00038716485,0.00006074216,0.00050032756,0.0021995318,0.0097893365,0.0023430178],"genre_scores_gemma":[0.057692938,0.00042338663,0.93166757,0.00029539116,0.00006892798,0.00042965103,0.007725142,0.000873579,0.0008234711],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959786,0.0012992881,0.00077148946,0.0006576549,0.0010801128,0.00021283902],"domain_scores_gemma":[0.9796606,0.015409562,0.0010148651,0.0017068179,0.0019870948,0.0002211766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058935573,0.0018950757,0.0015403617,0.0074228197,0.0015511299,0.003499798,0.0030162998,0.0015405898,0.0060318466],"category_scores_gemma":[0.018445142,0.0009938953,0.002196376,0.0053854007,0.0012751282,0.005296888,0.0021422012,0.0022409118,0.0035086565],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046718464,0.000630521,0.0069433427,0.002242606,0.0004806101,0.0016494391,0.001289416,0.037066985,0.030444873,0.09366968,0.049202047,0.77591336],"study_design_scores_gemma":[0.00025920465,0.00027443262,0.0044042463,0.00053036254,0.00071946176,0.0021435174,0.0008884523,0.57898915,0.078660995,0.24537706,0.087406546,0.00034644522],"about_ca_topic_score_codex":0.0036415723,"about_ca_topic_score_gemma":0.007037573,"teacher_disagreement_score":0.0074228197,"about_ca_system_score_codex":0.001283817,"about_ca_system_score_gemma":0.0030569418,"threshold_uncertainty_score":0.03116846},"labels":[],"label_agreement":null},{"id":"W2003601914","doi":"10.1186/1471-2105-13-s11-s7","title":"Biological event composition","year":2012,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Coreference; Biomedical text mining; Event (particle physics); Natural language processing; Task (project management); Artificial intelligence; Parsing; Information extraction; Resolution (logic); Information retrieval; Text mining","score_opus":0.037038019145623165,"score_gpt":0.29050173312038574,"score_spread":0.2534637139747626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2003601914","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019223448,0.0031819867,0.81951714,0.0024194191,0.0014720901,0.0017726934,0.021422029,0.041982863,0.08900823],"genre_scores_gemma":[0.12164637,0.002492975,0.74837387,0.0017702867,0.0009804828,0.0011088335,0.06608334,0.006235595,0.051308207],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962991,0.00047772794,0.0005056218,0.001419633,0.0010885213,0.00020949956],"domain_scores_gemma":[0.99540955,0.001676575,0.00032661628,0.00095096056,0.0013207037,0.0003156383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002791653,0.0021153179,0.0009919921,0.004507674,0.0025092822,0.0037032675,0.0025538139,0.001604769,0.052531287],"category_scores_gemma":[0.007792289,0.0007551303,0.0021883827,0.0029499526,0.0011232782,0.005725375,0.0051589967,0.0018536487,0.023854062],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014811446,0.00033417266,0.0078016785,0.00387417,0.00020839521,0.0026691244,0.0023466498,0.0047345846,0.04282884,0.11427283,0.113189496,0.7062589],"study_design_scores_gemma":[0.00009456349,0.0001255109,0.0033281907,0.00037577472,0.00024052506,0.0025797363,0.00061889907,0.027403766,0.052821185,0.10138873,0.8109006,0.00012257363],"about_ca_topic_score_codex":0.002744809,"about_ca_topic_score_gemma":0.0027359389,"teacher_disagreement_score":0.052531287,"about_ca_system_score_codex":0.0015332878,"about_ca_system_score_gemma":0.0026674226,"threshold_uncertainty_score":0.17573464},"labels":[],"label_agreement":null},{"id":"W2004441857","doi":"10.3138/jsp.44.3.005","title":"Source References and the Scientist's Mind-Map: Harvard vs. Vancouver Style","year":2013,"lang":"en","type":"article","venue":"Journal of Scholarly Publishing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Reading (process); Order (exchange); Computer science; Quality (philosophy); Space (punctuation); Ideal (ethics); Search engine indexing; Mind map; Process (computing); Information retrieval; World Wide Web; Artificial intelligence; Epistemology; Linguistics; Literature; Philosophy; Art","score_opus":0.013107228913788238,"score_gpt":0.2322029804361854,"score_spread":0.21909575152239716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004441857","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02950942,0.023699418,0.09617487,0.24732034,0.021817097,0.00021733294,0.0017336003,0.0012227129,0.5783052],"genre_scores_gemma":[0.606757,0.02712718,0.08383915,0.034008626,0.009691644,0.0005531071,0.002698852,0.002138493,0.23318605],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9873639,0.0057079983,0.0008769839,0.0009772701,0.0047685313,0.00030533198],"domain_scores_gemma":[0.9420178,0.026047288,0.0036149574,0.0062499987,0.016839724,0.005230238],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.014559676,0.00059316435,0.0005668811,0.009047899,0.0053402665,0.021270173,0.0018746993,0.0018550766,0.008759119],"category_scores_gemma":[0.06079474,0.00057476124,0.00036763528,0.012516557,0.016842773,0.008247588,0.0064620934,0.0034626666,0.0039530294],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009625926,0.000020623773,0.003027533,0.0003837507,0.00004423184,0.00018878613,0.022525616,0.00015835288,0.00061478885,0.6650266,0.21014023,0.097773306],"study_design_scores_gemma":[0.000028142362,0.000023680239,0.0037669453,0.00042662455,0.000030112033,0.00022512847,0.0054624663,0.00047685744,0.00042480588,0.21952757,0.7695636,0.000044145218],"about_ca_topic_score_codex":0.024644567,"about_ca_topic_score_gemma":0.04802716,"teacher_disagreement_score":0.97872984,"about_ca_system_score_codex":0.007626946,"about_ca_system_score_gemma":0.009205346,"threshold_uncertainty_score":0.07699984},"labels":[],"label_agreement":null},{"id":"W2004600423","doi":"10.3414/me14-01-0104","title":"Quality Assurance of UMLS Semantic Type Assignments Using SNOMED CT Hierarchies","year":2015,"lang":"en","type":"article","venue":"Methods of Information in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"New York Institute of Technology","funders":"U.S. National Library of Medicine; National Cancer Institute","keywords":"Unified Medical Language System; SNOMED CT; Computer science; Information retrieval; Hierarchy; Natural language processing; Quality assurance; Artificial intelligence; Quality (philosophy); Process (computing); Terminology; Medicine; Programming language; Linguistics; Pathology","score_opus":0.11900171250070093,"score_gpt":0.4593669719136684,"score_spread":0.34036525941296747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2004600423","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16071585,0.0007453985,0.82568645,0.0005195921,0.0001362498,0.0026309642,0.0019669365,0.004646752,0.0029517782],"genre_scores_gemma":[0.23396434,0.00020408009,0.7595245,0.00018881433,0.000031724794,0.0013428008,0.0033969574,0.0006926756,0.00065408816],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.90175146,0.051961783,0.013130042,0.007680951,0.024387814,0.0010879698],"domain_scores_gemma":[0.6463129,0.19216019,0.041256864,0.038526945,0.080435775,0.0013073726],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12571566,0.001291516,0.0010189663,0.016354084,0.002511282,0.004670094,0.0028452047,0.0016592307,0.0017396044],"category_scores_gemma":[0.24580523,0.0008932707,0.0018089009,0.005570635,0.0022524954,0.003265563,0.004056836,0.0014292037,0.00068048656],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021685245,0.0009722198,0.10235897,0.0043214075,0.0009405804,0.0021363527,0.03189118,0.046702944,0.07667401,0.028925024,0.012176583,0.6907323],"study_design_scores_gemma":[0.0004815426,0.0018074053,0.08522176,0.00579899,0.0015909536,0.0051710065,0.0076978183,0.44151163,0.2983097,0.047176983,0.10431657,0.0009156611],"about_ca_topic_score_codex":0.009314161,"about_ca_topic_score_gemma":0.010114,"teacher_disagreement_score":0.12571566,"about_ca_system_score_codex":0.0038652653,"about_ca_system_score_gemma":0.008297375,"threshold_uncertainty_score":0.6648559},"labels":[],"label_agreement":null},{"id":"W2005255873","doi":"10.1186/1471-2105-13-249","title":"Quantitative biomedical annotation using medical subject heading over-representation profiles (MeSHOPs)","year":2012,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Ontario Institute for Cancer Research; Child and Family Research Institute; University of British Columbia","funders":"National Institute of General Medical Sciences; Michael Smith Health Research BC; Canadian Institutes of Health Research; Government of Ontario; Ontario Institute for Cancer Research","keywords":"Computer science; Controlled vocabulary; Annotation; Information retrieval; Vocabulary; Ontology; Representation (politics); Subject (documents); Set (abstract data type); Unified Medical Language System; Open Biomedical Ontologies; Visualization; Data science; Natural language processing; Data mining; Artificial intelligence; World Wide Web; Ontology alignment","score_opus":0.06687942214877576,"score_gpt":0.36950989404859996,"score_spread":0.30263047189982417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2005255873","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38684264,0.005214366,0.48613676,0.00179646,0.00042013446,0.0015388272,0.08782367,0.011826316,0.018400803],"genre_scores_gemma":[0.58341676,0.00156488,0.37228045,0.00025540398,0.00027409315,0.0015434977,0.038164143,0.00079300965,0.0017077288],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961357,0.0011231304,0.00063081173,0.0007801522,0.0011946241,0.00013557718],"domain_scores_gemma":[0.9685819,0.018551156,0.0070342887,0.0018058647,0.0035303696,0.00049632805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049366476,0.0008310185,0.000526594,0.020400852,0.0005758585,0.002716861,0.0005932952,0.00065408106,0.0038750689],"category_scores_gemma":[0.03929156,0.00022612563,0.00080432725,0.017590737,0.0005121866,0.0021691096,0.0021484084,0.0006333671,0.0012573034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014731222,0.00021322707,0.1533525,0.0075052464,0.00080489944,0.0007127451,0.00421479,0.017171679,0.060120415,0.038207605,0.035500113,0.68072367],"study_design_scores_gemma":[0.00024414936,0.0010159435,0.34352642,0.0019069915,0.0011436472,0.0030465052,0.004781225,0.21259485,0.067704156,0.15010817,0.213346,0.0005819705],"about_ca_topic_score_codex":0.0015657322,"about_ca_topic_score_gemma":0.001622492,"teacher_disagreement_score":0.020400852,"about_ca_system_score_codex":0.0010147169,"about_ca_system_score_gemma":0.0010706751,"threshold_uncertainty_score":0.026107788},"labels":[],"label_agreement":null},{"id":"W2007691077","doi":"10.5339/qfarf.2013.biop-035","title":"A Semantic Web Framework To Computerize And Execute Clinical Guidelines: Towards The Handling Of Co-Morbidities In Clinical Decision Support Systems","year":2013,"lang":"en","type":"article","venue":"Qatar Foundation Annual Research Forum Volume 2013 Issue 1","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University","funders":"","keywords":"Clinical decision support system; Interoperability; Ontology; Decision support system; Computer science; Semantic interoperability; Medicine; Psychological intervention; Atrial fibrillation; Knowledge management; Medical emergency; Intensive care medicine; Data mining; Nursing; World Wide Web; Internal medicine","score_opus":0.13401818894898665,"score_gpt":0.4805223974655111,"score_spread":0.34650420851652447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007691077","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024311678,0.00029159503,0.98864555,0.00163205,0.00006277067,0.00046799655,0.00051192264,0.0028472021,0.0031097939],"genre_scores_gemma":[0.020525405,0.0004157482,0.9756049,0.00038911417,0.000037959137,0.00036308396,0.0014643697,0.0002001364,0.0009993169],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924522,0.0028739688,0.0014553089,0.000969748,0.0018526898,0.00039603355],"domain_scores_gemma":[0.99336845,0.0032984724,0.0005419537,0.0011508727,0.0011377315,0.00050240237],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012222936,0.0012608796,0.0012199694,0.0063000675,0.0022093346,0.008406539,0.0037988885,0.003758005,0.002400069],"category_scores_gemma":[0.012237908,0.0013156183,0.0042741117,0.0046977475,0.0038012734,0.008399332,0.0053994763,0.003860403,0.001484044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017511139,0.00057565566,0.0026031488,0.0012166036,0.00037504025,0.0016388851,0.00393541,0.09243769,0.0069328477,0.6593952,0.015087037,0.2156274],"study_design_scores_gemma":[0.00011203669,0.00010341659,0.00077590405,0.001209296,0.00018352622,0.0005103613,0.0013347176,0.4303192,0.0056529464,0.3637541,0.19592357,0.00012091154],"about_ca_topic_score_codex":0.027009144,"about_ca_topic_score_gemma":0.027482368,"teacher_disagreement_score":0.027009144,"about_ca_system_score_codex":0.0031753562,"about_ca_system_score_gemma":0.008817687,"threshold_uncertainty_score":0.06464183},"labels":[],"label_agreement":null},{"id":"W2007701698","doi":"10.1002/mar.10029","title":"Similarity of drug names: Comparison of objective and subjective measures","year":2002,"lang":"en","type":"article","venue":"Psychology and Marketing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University; University of Toronto; McGill University","funders":"","keywords":"Similarity (geometry); Multidimensional scaling; Psychology; Trigram; Correlation; Set (abstract data type); Statistics; Spelling; Variance (accounting); Mathematics; Natural language processing; Artificial intelligence; Linguistics; Computer science","score_opus":0.029787305261502202,"score_gpt":0.32274646466104834,"score_spread":0.29295915939954614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2007701698","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99207497,0.00020578745,0.0049495744,0.00005447871,0.000015020372,0.000045611876,0.00022175703,0.00002335789,0.0024094225],"genre_scores_gemma":[0.99800557,0.000046334433,0.001615022,0.00001055441,0.0000121759,0.00002181821,0.00016546744,0.000005069865,0.00011796882],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99487853,0.0022507238,0.0005923845,0.0004069753,0.0017829763,0.000088307],"domain_scores_gemma":[0.9330593,0.041282978,0.015417833,0.0031070916,0.005696517,0.0014363176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0047490634,0.00024283405,0.00031994944,0.0023430008,0.00028387987,0.0013483448,0.0002793848,0.00033294407,0.0018013271],"category_scores_gemma":[0.067532934,0.00013808267,0.00027860518,0.0014616207,0.00070753443,0.0011734975,0.0010689716,0.00035858498,0.0002016111],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009784319,0.00028350975,0.92221844,0.0003152024,0.0005897631,0.000107420834,0.0024892415,0.001508601,0.003974173,0.0015682443,0.00052189454,0.065445095],"study_design_scores_gemma":[0.000039030354,0.00070746435,0.9836706,0.000037778904,0.000079604615,0.0002761209,0.0018395075,0.008020082,0.001957312,0.002327836,0.0009865597,0.000058174424],"about_ca_topic_score_codex":0.0004237898,"about_ca_topic_score_gemma":0.0006166038,"teacher_disagreement_score":0.0047490634,"about_ca_system_score_codex":0.00037356533,"about_ca_system_score_gemma":0.00018883382,"threshold_uncertainty_score":0.025115728},"labels":[],"label_agreement":null},{"id":"W2008169899","doi":"10.2196/medinform.2671","title":"An Intelligent Content Discovery Technique for Health Portal Content Management","year":2014,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Australian Government","keywords":"Computer science; Content (measure theory); Content management; World Wide Web","score_opus":0.0393838111215607,"score_gpt":0.33906481233270797,"score_spread":0.29968100121114727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2008169899","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015822504,0.00012992424,0.976143,0.00028817722,0.00003390353,0.0005586259,0.00025822214,0.0045152875,0.002250432],"genre_scores_gemma":[0.0799768,0.00008939458,0.91618365,0.0000777116,0.000036136113,0.0002929076,0.00055876235,0.000170992,0.002613542],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.995146,0.001393962,0.00046819652,0.00076817616,0.0019930182,0.00023066251],"domain_scores_gemma":[0.9915235,0.00380349,0.000767582,0.001430486,0.0022614622,0.00021340406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034582021,0.000884418,0.00092796347,0.0075918273,0.0015326206,0.0025872986,0.0017186931,0.0013601551,0.0033385218],"category_scores_gemma":[0.010568335,0.0005348078,0.0013065602,0.005425806,0.001006695,0.0032892898,0.0021722557,0.0012033501,0.0022061781],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044976463,0.00064152246,0.0033212642,0.0004966874,0.000122341,0.0005037361,0.001855759,0.008091203,0.059544265,0.01985026,0.009083486,0.89603966],"study_design_scores_gemma":[0.00020173342,0.0006738732,0.004215227,0.00012259731,0.00028560017,0.0024194012,0.0013052699,0.76427126,0.13794197,0.027348392,0.061016977,0.00019776457],"about_ca_topic_score_codex":0.00236269,"about_ca_topic_score_gemma":0.0028845668,"teacher_disagreement_score":0.0075918273,"about_ca_system_score_codex":0.0014157507,"about_ca_system_score_gemma":0.0018974554,"threshold_uncertainty_score":0.01828891},"labels":[],"label_agreement":null},{"id":"W2011175967","doi":"10.1108/00220410810867551","title":"Classification, interdisciplinarity, and the study of science","year":2008,"lang":"en","type":"article","venue":"Journal of Documentation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":60,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Originality; Computer science; Epistemology; Value (mathematics); Scholarship; Point (geometry); Sociology; Data science; Management science; Social science; Mathematics; Machine learning; Philosophy; Political science; Law","score_opus":0.027149219746073965,"score_gpt":0.3443138993624349,"score_spread":0.3171646796163609,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011175967","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17286186,0.14170325,0.20300001,0.13140047,0.0032097735,0.0007955403,0.0005138452,0.00028781418,0.3462274],"genre_scores_gemma":[0.93288153,0.014065141,0.045262367,0.003074932,0.0007243215,0.00059038354,0.000240478,0.000054183172,0.0031067368],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.93918973,0.045888927,0.0035415401,0.0030641158,0.007166658,0.001148975],"domain_scores_gemma":[0.88002753,0.08698754,0.011632659,0.011823906,0.007368252,0.0021601154],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.040071957,0.00053874275,0.0013107712,0.018917058,0.008605632,0.016940981,0.0019881527,0.0023139513,0.0028286292],"category_scores_gemma":[0.05388379,0.00038438276,0.0007256818,0.026913922,0.05182801,0.022012612,0.007403031,0.003030535,0.0004148119],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014980485,0.000032225456,0.0031147662,0.0005012181,0.000029609977,0.00008021623,0.025273232,0.00026766016,0.00006849627,0.9272568,0.0020141238,0.04134656],"study_design_scores_gemma":[0.000010093074,0.000012498312,0.0018725019,0.0010055141,0.000013769978,0.00011673191,0.019349998,0.00055466854,0.0000903516,0.94008,0.036872156,0.000021585058],"about_ca_topic_score_codex":0.0059251348,"about_ca_topic_score_gemma":0.0036128603,"teacher_disagreement_score":0.99139434,"about_ca_system_score_codex":0.010420123,"about_ca_system_score_gemma":0.010920329,"threshold_uncertainty_score":0.21192324},"labels":[],"label_agreement":null},{"id":"W2011628373","doi":"10.3414/me12-01-0029","title":"Formative Evaluation of Ontology Learning Methods for Entity Discovery by Using Existing Ontologies as Reference Standards","year":2013,"lang":"en","type":"article","venue":"Methods of Information in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"National Institutes of Health; National Cancer Institute; University of Pittsburgh; Radiological Society of North America","keywords":"Ontology; Computer science; Formative assessment; Domain (mathematical analysis); Set (abstract data type); Information retrieval; Data mining; Natural language processing; Artificial intelligence; Statistics; Mathematics","score_opus":0.11936021319266922,"score_gpt":0.5145004667197498,"score_spread":0.3951402535270806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2011628373","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16209094,0.0016120694,0.8213636,0.0006261491,0.00020966872,0.00096703414,0.00083576614,0.006915469,0.0053792614],"genre_scores_gemma":[0.45179984,0.0002707652,0.54315525,0.0001684629,0.00007309,0.001067722,0.0018574717,0.00054731575,0.0010600197],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.90518886,0.048734084,0.007840811,0.006178717,0.030818854,0.0012386695],"domain_scores_gemma":[0.69554454,0.20089278,0.014775049,0.024491187,0.06232876,0.0019677428],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08058046,0.0023282818,0.0014642257,0.008676305,0.0014701735,0.0052340673,0.0035554306,0.0021386202,0.0016093104],"category_scores_gemma":[0.23998877,0.0005782474,0.0014923213,0.0040248153,0.0016168647,0.005583088,0.004088328,0.002236101,0.00077294983],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023187285,0.0014604157,0.04139418,0.0012906967,0.0012677934,0.00029651108,0.0016324992,0.114133134,0.017752169,0.010508787,0.0062965085,0.8016485],"study_design_scores_gemma":[0.00029429776,0.0015197207,0.013869392,0.00027096883,0.00022532658,0.00032745296,0.00089006696,0.9069944,0.06147753,0.0070285243,0.0069105322,0.0001917362],"about_ca_topic_score_codex":0.0046014097,"about_ca_topic_score_gemma":0.0056516137,"teacher_disagreement_score":0.08058046,"about_ca_system_score_codex":0.0046706726,"about_ca_system_score_gemma":0.0035576695,"threshold_uncertainty_score":0.42615527},"labels":[],"label_agreement":null},{"id":"W2016928406","doi":"10.1093/bioinformatics/bth451","title":"Discovering patterns to extract protein–protein interactions from full texts","year":2004,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":238,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Recall rate; Scientific literature; Matching (statistics); Information retrieval; Natural language processing; Artificial intelligence; Data mining; Biology","score_opus":0.015499909559941955,"score_gpt":0.267165727324153,"score_spread":0.25166581776421104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2016928406","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14015535,0.0022224074,0.82563525,0.001515233,0.00011106002,0.0008665443,0.012340663,0.01188825,0.0052652713],"genre_scores_gemma":[0.14158413,0.0009543115,0.8376223,0.00018710073,0.00011106727,0.00047838074,0.016636407,0.00033595064,0.0020904369],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983241,0.00029749412,0.00032063996,0.00056244235,0.00041349788,0.000081860446],"domain_scores_gemma":[0.99323684,0.0042275465,0.0008915794,0.0005086994,0.0009319963,0.00020336945],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016072782,0.001359032,0.00095645466,0.008935859,0.0007376109,0.0016514097,0.0013687036,0.001199044,0.0033064066],"category_scores_gemma":[0.008883455,0.0005236529,0.0010605433,0.0058569727,0.0006311505,0.0036376158,0.0014113077,0.0009292141,0.0035135765],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006136655,0.00067552633,0.023975987,0.002244829,0.00028940203,0.0014783839,0.0008317591,0.005176969,0.06321051,0.0062115504,0.015484469,0.879807],"study_design_scores_gemma":[0.00042050757,0.00086425524,0.051733825,0.0006936747,0.0010279241,0.010380289,0.002551163,0.5369634,0.2139628,0.09066459,0.090519115,0.00021845313],"about_ca_topic_score_codex":0.0009805065,"about_ca_topic_score_gemma":0.0015919708,"teacher_disagreement_score":0.008935859,"about_ca_system_score_codex":0.00044651856,"about_ca_system_score_gemma":0.0010888096,"threshold_uncertainty_score":0.011061072},"labels":[],"label_agreement":null},{"id":"W2017882693","doi":"10.1016/j.artmed.2014.03.005","title":"A token centric part-of-speech tagger for biomedical text","year":2014,"lang":"en","type":"article","venue":"Artificial Intelligence in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Artificial intelligence; Security token; Natural language processing; Classifier (UML); Part-of-speech tagging; Lemmatisation; Domain (mathematical analysis); Training set; Speech recognition; Part of speech","score_opus":0.05250969537362214,"score_gpt":0.34762459100350845,"score_spread":0.2951148956298863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2017882693","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035076022,0.002114368,0.8542795,0.0009094394,0.001039405,0.00052828377,0.045942802,0.05575794,0.0043521984],"genre_scores_gemma":[0.2011263,0.00096388604,0.72247255,0.00097573595,0.0004316372,0.00074399106,0.0617603,0.0021629427,0.009362685],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985917,0.00021661166,0.00023307065,0.0005077828,0.00033207182,0.0001186488],"domain_scores_gemma":[0.9947378,0.0026515536,0.00063970336,0.0006116376,0.0010642812,0.00029492893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018979603,0.0010872388,0.0012979457,0.0045263753,0.000762234,0.001644989,0.001469867,0.0020107592,0.0066060266],"category_scores_gemma":[0.005510037,0.0005284051,0.0010604863,0.003236014,0.00053591665,0.0024775865,0.0015991932,0.0014968514,0.012848582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021754578,0.00041115243,0.0063274903,0.0018156766,0.00034078985,0.0015046188,0.00048078256,0.0066053225,0.21662158,0.009154884,0.054254316,0.70030797],"study_design_scores_gemma":[0.00042422372,0.0009533917,0.020873617,0.0006057699,0.0010657936,0.005038524,0.00089910306,0.34713006,0.37988952,0.067948066,0.1747529,0.00041888887],"about_ca_topic_score_codex":0.0017420754,"about_ca_topic_score_gemma":0.00332221,"teacher_disagreement_score":0.0066060266,"about_ca_system_score_codex":0.00075394707,"about_ca_system_score_gemma":0.0024401266,"threshold_uncertainty_score":0.022099316},"labels":[],"label_agreement":null},{"id":"W2019106434","doi":"10.1016/s0277-9536(01)00246-5","title":"Making health data maps: a case study of a community/university research collaboration","year":2002,"lang":"en","type":"article","venue":"Social Science & Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":46,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Regent Park Community Health Centre; University of Toronto","funders":"","keywords":"Community health; Knowledge management; Information system; Process (computing); Action research; Computer science; Sociology; Medicine; Public health; Engineering; Nursing","score_opus":0.33951141061505924,"score_gpt":0.5028294475304497,"score_spread":0.16331803691539043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2019106434","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82191086,0.0009022925,0.10636782,0.035154875,0.0002380079,0.0012933813,0.0018214307,0.0013396674,0.03097172],"genre_scores_gemma":[0.88656396,0.000461073,0.10648731,0.00071351527,0.00007720234,0.00030489164,0.0008491514,0.00026476718,0.0042781755],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.97554135,0.018147679,0.0010860392,0.0015699965,0.0026391447,0.0010157513],"domain_scores_gemma":[0.8659261,0.106687635,0.0044954717,0.010015364,0.006034956,0.0068405014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026003899,0.00064957066,0.0005517595,0.0043735784,0.012846188,0.0067845094,0.004149214,0.007215873,0.0054839277],"category_scores_gemma":[0.06525195,0.000714932,0.0011446897,0.008955501,0.0040248553,0.0082783485,0.010125401,0.0029953383,0.0011019707],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00140945,0.0039567393,0.079336226,0.002144237,0.0004811688,0.048526146,0.39023665,0.022022467,0.0075920904,0.069026604,0.03636003,0.33890817],"study_design_scores_gemma":[0.000418519,0.0008527928,0.019891907,0.0008499016,0.0003777538,0.013780478,0.51570684,0.049330287,0.012228095,0.08584768,0.30034137,0.0003742935],"about_ca_topic_score_codex":0.016026018,"about_ca_topic_score_gemma":0.021955814,"teacher_disagreement_score":0.026003899,"about_ca_system_score_codex":0.002984446,"about_ca_system_score_gemma":0.0081952205,"threshold_uncertainty_score":0.13752341},"labels":[],"label_agreement":null},{"id":"W2020726632","doi":"10.1108/00220411111145061","title":"The modernity of classification","year":2011,"lang":"en","type":"article","venue":"Journal of Documentation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Modernity; Originality; Foundation (evidence); Mainstream; Epistemology; Sociology; Value (mathematics); Work (physics); Computer science; Engineering ethics; Data science; Social science; Philosophy; Law; Political science; Engineering; Machine learning","score_opus":0.052757113600748105,"score_gpt":0.3107569152511119,"score_spread":0.2579998016503638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2020726632","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046484787,0.04476118,0.46102884,0.20586523,0.0042925584,0.00024286247,0.0004080231,0.0006057344,0.23631082],"genre_scores_gemma":[0.82865465,0.012592923,0.12951341,0.010452363,0.004372157,0.0004560017,0.00030312655,0.00033738802,0.013318001],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.96490234,0.018263146,0.0021198294,0.0049903174,0.008429079,0.0012953178],"domain_scores_gemma":[0.9166996,0.04712456,0.003954273,0.01772166,0.012271743,0.0022281597],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.037377432,0.0007983164,0.0013284138,0.010169389,0.009976003,0.017864134,0.0025731677,0.0034810333,0.0051166047],"category_scores_gemma":[0.046626333,0.0006202825,0.0010072049,0.0063400296,0.0895966,0.031527802,0.008771672,0.008200912,0.0014613875],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009836241,0.0000074230425,0.000499195,0.00008065943,0.0000073308543,0.000014514939,0.0022730285,0.000134348,0.000056419016,0.9787048,0.0022284023,0.01598398],"study_design_scores_gemma":[0.0000057008965,0.000013821659,0.00038170634,0.00024524934,0.000006015515,0.00007685282,0.0013815552,0.00071846973,0.00012466784,0.93679106,0.060237307,0.000017507993],"about_ca_topic_score_codex":0.005022291,"about_ca_topic_score_gemma":0.002476537,"teacher_disagreement_score":0.990024,"about_ca_system_score_codex":0.013297655,"about_ca_system_score_gemma":0.008709499,"threshold_uncertainty_score":0.19767314},"labels":[],"label_agreement":null},{"id":"W2021074129","doi":"10.1371/journal.pone.0154556","title":"The Ontology for Biomedical Investigations","year":2016,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":385,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Biotechnology and Biological Sciences Research Council; National Institute of Diabetes and Digestive and Kidney Diseases; National Institute of Allergy and Infectious Diseases; Division of Biological Infrastructure; National Institute of General Medical Sciences; National Center for Advancing Translational Sciences; Agence Nationale de la Recherche; National Human Genome Research Institute; Engineering and Physical Sciences Research Council; National Institute on Drug Abuse; California Institute for Regenerative Medicine; National Institutes of Health; National Science Foundation","keywords":"Open Biomedical Ontologies; Ontology; Computer science; Process ontology; Data science; Upper ontology; Interoperability; Biological database; Ontology-based data integration; World Wide Web; Information retrieval; Semantic Web; Suggested Upper Merged Ontology; Bioinformatics; Biology","score_opus":0.060346101899732336,"score_gpt":0.2676917707597481,"score_spread":0.20734566886001576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2021074129","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002370153,0.034460317,0.4543953,0.08879756,0.008306142,0.0011741621,0.009603845,0.004453502,0.39643905],"genre_scores_gemma":[0.07069177,0.04936327,0.7131243,0.040306516,0.007542742,0.0029676536,0.02262923,0.0019887968,0.09138574],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9845879,0.0057545733,0.002506776,0.0018253644,0.0043554464,0.0009699371],"domain_scores_gemma":[0.9798237,0.009188528,0.0014884506,0.004265044,0.0038025286,0.0014316809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012873457,0.0011708819,0.0016754367,0.008848258,0.0050262697,0.0149798365,0.0035491486,0.007353626,0.011616969],"category_scores_gemma":[0.02067108,0.0010996384,0.002950602,0.010996636,0.011263593,0.02117115,0.009409229,0.010841732,0.008312205],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000119225915,0.00001806461,0.00014976625,0.0003673799,0.00001827253,0.00012612712,0.0008311861,0.00027280743,0.0002961986,0.9149504,0.0456282,0.037329648],"study_design_scores_gemma":[0.000009057319,0.0000056691742,0.00019775507,0.0004446465,0.000011181588,0.00029940595,0.00023308798,0.0004547499,0.000126872,0.2357373,0.7624568,0.000023539344],"about_ca_topic_score_codex":0.01308201,"about_ca_topic_score_gemma":0.007374669,"teacher_disagreement_score":0.0149798365,"about_ca_system_score_codex":0.007992084,"about_ca_system_score_gemma":0.020402797,"threshold_uncertainty_score":0.068082154},"labels":[],"label_agreement":null},{"id":"W2022602119","doi":"10.1210/me.2015-1119","title":"Editorial: What's in a Name, or the Impact of Misnomers in Endocrine Research","year":2015,"lang":"en","type":"editorial","venue":"Molecular Endocrinology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"McGill University","keywords":"Misnomer; Object (grammar); Disease; Misattribution of memory; Epistemology; Psychology; Biology; Medicine; Psychiatry; Computer science; Pathology; Cognition; Philosophy","score_opus":0.03945865314673591,"score_gpt":0.40969407613632336,"score_spread":0.37023542298958745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022602119","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007497627,0.0036654596,0.000111914145,0.01883809,0.9759857,0.000048929174,0.0001204316,0.00011812692,0.0010363775],"genre_scores_gemma":[0.00065606483,0.0036932358,0.00014184337,0.02000234,0.9704184,0.000041403284,0.000068469606,0.00007019574,0.004908032],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99032307,0.0015880989,0.0014667413,0.0013668841,0.00465184,0.000603277],"domain_scores_gemma":[0.94925404,0.020202333,0.004928714,0.0013717855,0.019716041,0.004527063],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010211546,0.0041142744,0.005141789,0.0053867744,0.003338287,0.009686537,0.0055337497,0.013760849,0.04234439],"category_scores_gemma":[0.059292622,0.0013828995,0.0034190267,0.0022104122,0.0033750962,0.005804001,0.0014471111,0.010765239,0.021632614],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000102352875,0.000033192788,0.00008533751,0.0009972572,0.00005593097,0.00012366971,0.00002757579,0.000044344903,0.00009131625,0.00028979534,0.9918311,0.0063181063],"study_design_scores_gemma":[0.00021992123,0.000174011,0.0006991352,0.0023185485,0.00017898508,0.0006533617,0.00011546286,0.00034317272,0.00038334922,0.0013433808,0.9935068,0.000063865664],"about_ca_topic_score_codex":0.0011949582,"about_ca_topic_score_gemma":0.0019288154,"teacher_disagreement_score":0.04234439,"about_ca_system_score_codex":0.0038528114,"about_ca_system_score_gemma":0.004504752,"threshold_uncertainty_score":0.1416561},"labels":[],"label_agreement":null},{"id":"W2022692405","doi":"10.1186/1755-8794-6-s2-s3","title":"Compensating for literature annotation bias when predicting novel drug-disease relationships through Medical Subject Heading Over-representation Profile (MeSHOP) similarity","year":2013,"lang":"en","type":"article","venue":"BMC Medical Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Ontario Institute for Cancer Research; Child and Family Research Institute; University of British Columbia","funders":"National Institute of General Medical Sciences; Michael Smith Health Research BC; Canadian Institutes of Health Research; Government of Ontario; Ontario Institute for Cancer Research","keywords":"Leverage (statistics); Disease; Annotation; MEDLINE; Computer science; Precision medicine; Subject (documents); Representation (politics); Medicine; Information retrieval; Bioinformatics; Artificial intelligence; Biology; Pathology; World Wide Web","score_opus":0.08568723756780985,"score_gpt":0.32623566111691327,"score_spread":0.2405484235491034,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2022692405","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7844471,0.012262526,0.15985739,0.0018150884,0.0004883493,0.00088740536,0.028979918,0.0051384065,0.0061239526],"genre_scores_gemma":[0.84651875,0.0021456378,0.11983501,0.00023404308,0.00044934108,0.0005564965,0.029268548,0.0002449095,0.0007472813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9930354,0.0023649833,0.0013627312,0.0013974217,0.0015624046,0.000276998],"domain_scores_gemma":[0.94940346,0.035707545,0.007409801,0.0028345906,0.0038551018,0.00078963774],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010417527,0.0012629469,0.0014029278,0.02555136,0.0009582712,0.0028556203,0.0010803291,0.001290206,0.0015693912],"category_scores_gemma":[0.06239062,0.0002513485,0.001292591,0.018568581,0.00052658195,0.0027731115,0.002781036,0.0009815987,0.0009985518],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00177152,0.0004959011,0.44352347,0.005247194,0.0019225178,0.0012513401,0.00090876454,0.020714788,0.029211378,0.0038208147,0.015980784,0.47515163],"study_design_scores_gemma":[0.00041707297,0.0014820297,0.40034825,0.0012158074,0.0034567348,0.0044519063,0.0018938545,0.46940142,0.037613917,0.037713315,0.041714836,0.00029087605],"about_ca_topic_score_codex":0.0028199055,"about_ca_topic_score_gemma":0.004075325,"teacher_disagreement_score":0.02555136,"about_ca_system_score_codex":0.00087840477,"about_ca_system_score_gemma":0.0024156268,"threshold_uncertainty_score":0.055093765},"labels":[],"label_agreement":null},{"id":"W2023917861","doi":"10.1038/npre.2010.5060.1","title":"Bio2RDF: Convert, Provide And Reuse.","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Université Laval","funders":"","keywords":"Computer science; Identifier; SPARQL; Linked data; RDF; Information retrieval; World Wide Web; Semantic Web; Database; Data science","score_opus":0.0077923154753697,"score_gpt":0.27671304805933983,"score_spread":0.26892073258397015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2023917861","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020042835,0.0009279183,0.13947284,0.0009827286,0.00075737445,0.0004232923,0.11523768,0.71472436,0.025469542],"genre_scores_gemma":[0.01932269,0.0014034051,0.11848218,0.0014716364,0.00017825303,0.0011179481,0.48701617,0.34720436,0.023803405],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9936527,0.0007177136,0.0006735393,0.0011695484,0.0030618317,0.00072467566],"domain_scores_gemma":[0.99155784,0.0017951366,0.00043430977,0.004289203,0.0014712982,0.00045231712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009089391,0.0039450824,0.0017116641,0.007727768,0.0017230485,0.007955361,0.0050044116,0.0031973212,0.09437542],"category_scores_gemma":[0.02307089,0.002721341,0.003670993,0.006320156,0.0015660737,0.010426536,0.009724925,0.004104895,0.1353785],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00073210744,0.00027686407,0.002029183,0.0013478288,0.00017782749,0.00084447674,0.00068020506,0.0014424275,0.0052457335,0.024250995,0.8437938,0.11917862],"study_design_scores_gemma":[0.0001215838,0.000042280684,0.001505758,0.00032939896,0.00004444304,0.0005774161,0.00016172668,0.0026141745,0.011668183,0.0158392,0.96693635,0.00015951983],"about_ca_topic_score_codex":0.0071131466,"about_ca_topic_score_gemma":0.00451961,"teacher_disagreement_score":0.09437542,"about_ca_system_score_codex":0.0029504956,"about_ca_system_score_gemma":0.0033418727,"threshold_uncertainty_score":0.31571722},"labels":[],"label_agreement":null},{"id":"W2024038403","doi":"10.1093/nar/gku1068","title":"BRENDA in 2015: exciting developments in its 25th year of existence","year":2014,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":195,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Bundesministerium für Bildung und Forschung","keywords":"Annotation; Information retrieval; Relevance (law); Computer science; KEGG; Computational biology; Function (biology); Ontology; Biology; Data mining; Artificial intelligence; Gene ontology; Biochemistry; Genetics; Gene","score_opus":0.06924922146140024,"score_gpt":0.38530854474814724,"score_spread":0.316059323286747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024038403","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006843416,0.26616335,0.08694007,0.117884904,0.2792604,0.0003583991,0.0140478695,0.031190794,0.1973108],"genre_scores_gemma":[0.03235179,0.12797374,0.08033311,0.025252907,0.0487703,0.00049864064,0.033303693,0.016239164,0.6352766],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9954573,0.00071622076,0.00034524425,0.0007995085,0.001957937,0.0007238345],"domain_scores_gemma":[0.98089534,0.0018380011,0.000688002,0.0014419387,0.004279888,0.01085679],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01734525,0.0023371908,0.00207909,0.004446201,0.0016328986,0.010926501,0.0037321814,0.0036620086,0.071884826],"category_scores_gemma":[0.013753675,0.0008790424,0.0014665988,0.0037429037,0.0019352563,0.008882305,0.0096138595,0.004986664,0.08275592],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029888164,0.00008134651,0.00074715883,0.0004327169,0.000024391995,0.00020663689,0.00023556584,0.00021102498,0.0014681601,0.012616263,0.7281356,0.25554228],"study_design_scores_gemma":[0.0000067973433,0.00002702246,0.00022888032,0.00009153881,0.0000032282285,0.00007450183,0.000033435048,0.000056454726,0.00031509664,0.00091198157,0.99823594,0.000015107945],"about_ca_topic_score_codex":0.003861388,"about_ca_topic_score_gemma":0.0033993253,"teacher_disagreement_score":0.071884826,"about_ca_system_score_codex":0.0037056224,"about_ca_system_score_gemma":0.007348139,"threshold_uncertainty_score":0.2404787},"labels":[],"label_agreement":null},{"id":"W2024912328","doi":"10.1016/j.jbi.2009.12.005","title":"An ontological modeling approach to cerebrovascular disease studies: The NEUROWEB case","year":2010,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institutes of Health","keywords":"Exploit; Computer science; Ontology; Phenotype; Data mining; Information retrieval; Clinical phenotype; Association (psychology); Robustness (evolution); Biology; Psychology","score_opus":0.04029101152617153,"score_gpt":0.33008018202268885,"score_spread":0.2897891704965173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2024912328","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17454425,0.0037697328,0.725418,0.03292684,0.00023147932,0.0005673728,0.0040614037,0.0005770929,0.05790394],"genre_scores_gemma":[0.7059248,0.0022727908,0.28515282,0.0015660415,0.00014133062,0.0002431037,0.0014243237,0.0000991723,0.003175704],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99755883,0.0014863354,0.0002481937,0.0002036223,0.00037945926,0.00012351118],"domain_scores_gemma":[0.9929416,0.0052933698,0.0004218502,0.0006265937,0.0003938905,0.00032259047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054395306,0.00044413644,0.00036638812,0.004182939,0.0020681333,0.00407369,0.001706581,0.0020708346,0.0016482241],"category_scores_gemma":[0.008692389,0.00033950905,0.0009833196,0.0045498614,0.0019187214,0.0045806775,0.002830688,0.0014228494,0.00022279333],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032579296,0.00047426173,0.03406851,0.0005286001,0.0003380683,0.013822952,0.0071742255,0.034051646,0.0043521766,0.76547235,0.010797802,0.12859361],"study_design_scores_gemma":[0.00008683049,0.00007139757,0.0068502785,0.0004914541,0.0004195876,0.007941345,0.0074838,0.19193557,0.0034722292,0.689911,0.09125289,0.000083594416],"about_ca_topic_score_codex":0.020741574,"about_ca_topic_score_gemma":0.029603213,"teacher_disagreement_score":0.020741574,"about_ca_system_score_codex":0.0026788777,"about_ca_system_score_gemma":0.0037404296,"threshold_uncertainty_score":0.041241705},"labels":[],"label_agreement":null},{"id":"W2025189988","doi":"10.1109/jbhi.2014.2383840","title":"Exploiting Semantic Web Technologies to Develop OWL-Based Clinical Practice Guideline Execution Engines","year":2014,"lang":"en","type":"article","venue":"IEEE Journal of Biomedical and Health Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Semantic Web Rule Language; Semantic Web; OWL-S; Web Ontology Language; Semantic reasoner; SPARQL; World Wide Web; Ontology; Semantic Web Stack; Correctness; Information retrieval; Database; Programming language; Semantic analytics; RDF; Artificial intelligence","score_opus":0.04392282952750862,"score_gpt":0.3851935661613002,"score_spread":0.3412707366337916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025189988","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055654077,0.00019753561,0.9691177,0.00059400813,0.00010733784,0.0014464734,0.0015414093,0.015922602,0.005507629],"genre_scores_gemma":[0.03298444,0.0004148139,0.958259,0.00046429338,0.00002418333,0.0005187147,0.0043045324,0.0006873391,0.0023425857],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968868,0.00067412853,0.00072274654,0.00035586316,0.0011835584,0.0001769684],"domain_scores_gemma":[0.99551827,0.001982958,0.00039382663,0.0006576601,0.0013001779,0.00014705761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004547178,0.0009997874,0.0006142717,0.0023013526,0.00056296954,0.002977789,0.0018058677,0.0011528883,0.0035065345],"category_scores_gemma":[0.01143528,0.00076908554,0.0017620005,0.0015298015,0.0007262826,0.0024539009,0.0015580107,0.0018989542,0.0021161917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043958347,0.0011637687,0.0055418448,0.0024753388,0.0005306646,0.002348304,0.0015593455,0.10595491,0.033246003,0.14231277,0.038354345,0.66607314],"study_design_scores_gemma":[0.00032176086,0.0002137954,0.001771604,0.0007983699,0.00028671068,0.00086817745,0.00048518574,0.6868181,0.055415638,0.07739923,0.17544486,0.00017654184],"about_ca_topic_score_codex":0.016964654,"about_ca_topic_score_gemma":0.0161974,"teacher_disagreement_score":0.016964654,"about_ca_system_score_codex":0.00159212,"about_ca_system_score_gemma":0.004597072,"threshold_uncertainty_score":0.03373182},"labels":[],"label_agreement":null},{"id":"W2025526593","doi":"10.1038/npre.2009.3570.1","title":"Open Biomedical Ontologies Applied to Prostate Cancer","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"London Health Sciences Centre; Western University","funders":"London Health Sciences Centre","keywords":"SNOMED CT; Ontology; Computer science; Open Biomedical Ontologies; Controlled vocabulary; Information retrieval; DICOM; World Wide Web; Terminology; Upper ontology; Semantic Web; Artificial intelligence; Suggested Upper Merged Ontology","score_opus":0.017104730573146845,"score_gpt":0.3500854965580411,"score_spread":0.3329807659848943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025526593","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021871872,0.006442423,0.92592806,0.007597294,0.000531899,0.00047909893,0.006768886,0.0039190124,0.026461396],"genre_scores_gemma":[0.17822449,0.0052210214,0.79772586,0.0014657506,0.00027470913,0.00038594299,0.0113358,0.0007208027,0.0046456507],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99010485,0.004073042,0.0013209927,0.0011192474,0.0030780022,0.00030383194],"domain_scores_gemma":[0.9824258,0.0100274,0.0012755911,0.0033887592,0.0024129048,0.00046946094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008190911,0.00062546524,0.000886711,0.014992943,0.0022896468,0.006931351,0.0016682731,0.0011901305,0.004062892],"category_scores_gemma":[0.03217636,0.0006032466,0.0016747141,0.019831385,0.002432633,0.006663207,0.0057468996,0.0016835275,0.0010391987],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010198607,0.000121720754,0.0058437977,0.0021824893,0.00019123868,0.0011718337,0.004318803,0.012322215,0.0037809454,0.58841974,0.021649221,0.35989594],"study_design_scores_gemma":[0.000034757482,0.000024598654,0.0047892225,0.0013516919,0.00014373241,0.00079683115,0.002162501,0.056309383,0.008239035,0.52998036,0.3960774,0.000090533664],"about_ca_topic_score_codex":0.016059319,"about_ca_topic_score_gemma":0.0121074505,"teacher_disagreement_score":0.016059319,"about_ca_system_score_codex":0.004173844,"about_ca_system_score_gemma":0.005708118,"threshold_uncertainty_score":0.043318212},"labels":[],"label_agreement":null},{"id":"W2025891673","doi":"10.1038/461171a","title":"Post-publication sharing of data and tools","year":2009,"lang":"en","type":"article","venue":"Nature","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":171,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"National Heart, Lung, and Blood Institute; U.S. National Library of Medicine; Wellcome Trust","keywords":"Data sharing; Data science; Computer science; World Wide Web; Medicine; Alternative medicine","score_opus":0.025650042291687658,"score_gpt":0.3222853500815559,"score_spread":0.29663530778986824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2025891673","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036144897,0.009226072,0.29038814,0.08317708,0.088036396,0.009048649,0.1513087,0.040078726,0.2925913],"genre_scores_gemma":[0.14964983,0.009970915,0.18118277,0.010510039,0.02094808,0.009065672,0.29119268,0.02402927,0.30345082],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.91246784,0.022598512,0.020675056,0.0070154546,0.032390073,0.004853132],"domain_scores_gemma":[0.39425716,0.12382072,0.020593314,0.32955724,0.11633789,0.015433684],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.114282005,0.0015519563,0.0031000064,0.023324277,0.0050087795,0.02141038,0.005137052,0.0046489784,0.08177491],"category_scores_gemma":[0.28113306,0.0018346335,0.003124902,0.030766476,0.0036343026,0.016967852,0.017093055,0.0068171523,0.0985612],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016444593,0.0003967575,0.007422589,0.0035442817,0.0006980086,0.0014428715,0.004301738,0.00079881516,0.014494934,0.02839311,0.5520969,0.3847656],"study_design_scores_gemma":[0.00015342417,0.00017225533,0.0069595356,0.0005819395,0.00020945536,0.0005544087,0.001101782,0.0006927129,0.010011687,0.019015074,0.9603716,0.00017612545],"about_ca_topic_score_codex":0.002036741,"about_ca_topic_score_gemma":0.0023642527,"teacher_disagreement_score":0.994863,"about_ca_system_score_codex":0.0028170217,"about_ca_system_score_gemma":0.01885906,"threshold_uncertainty_score":0.60438824},"labels":[],"label_agreement":null},{"id":"W2028968279","doi":"10.6026/97320630001360","title":"xGENIA: A comprehensive OWL ontology based on the GENIA corpus","year":2007,"lang":"en","type":"article","venue":"Bioinformation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Ontology; Annotation; Natural language processing; Information retrieval; Open Biomedical Ontologies; Taxonomy (biology); Benchmark (surveying); Unified Medical Language System; Artificial intelligence; Ontology alignment; Semantic Web; Ontology-based data integration","score_opus":0.020005021520630172,"score_gpt":0.26187743423711823,"score_spread":0.24187241271648807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2028968279","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041780073,0.0027086108,0.74777347,0.005701148,0.00067667465,0.0018542,0.1251775,0.027179029,0.047149263],"genre_scores_gemma":[0.11849338,0.0032707436,0.70467937,0.0016330832,0.00014957036,0.0016580753,0.15837182,0.0023191187,0.009424795],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993451,0.00013192209,0.0001364543,0.00013418932,0.00020345324,0.000048823975],"domain_scores_gemma":[0.99857354,0.00057860767,0.00018649116,0.00021745672,0.00034080457,0.00010300258],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015056654,0.00061623816,0.0005060775,0.0052990867,0.0013362571,0.001964064,0.0013104616,0.0007455898,0.0053928606],"category_scores_gemma":[0.0035718342,0.00043305257,0.00087572384,0.0049701408,0.00091297727,0.0036591059,0.0017898132,0.0009963494,0.0019449458],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046694995,0.00028716252,0.009507753,0.0049035745,0.0002425581,0.0028523274,0.004193509,0.0107469745,0.025100881,0.23534834,0.2538431,0.45250684],"study_design_scores_gemma":[0.00008818134,0.00006851059,0.010828018,0.0011618932,0.00015738535,0.0017670715,0.0013934411,0.034712017,0.007008924,0.04429654,0.89841896,0.0000990019],"about_ca_topic_score_codex":0.01577945,"about_ca_topic_score_gemma":0.016658407,"teacher_disagreement_score":0.01577945,"about_ca_system_score_codex":0.0016732847,"about_ca_system_score_gemma":0.00390602,"threshold_uncertainty_score":0.03137517},"labels":[],"label_agreement":null},{"id":"W2029921227","doi":"10.1093/bioinformatics/btl648","title":"DataBiNS: a BioMoby-based data-mining workflow for biological pathways and non-synonymous SNPs","year":2007,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital; University of British Columbia","funders":"Genome Canada","keywords":"Workflow; Computer science; Identifier; Web service; Workflow management system; World Wide Web; XML; Data mining; Information retrieval; Database; Programming language","score_opus":0.07184932710687342,"score_gpt":0.306166776981287,"score_spread":0.2343174498744136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2029921227","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018247329,0.0009802737,0.61933535,0.0018198857,0.00026356635,0.0010191095,0.15233162,0.19976021,0.0062426245],"genre_scores_gemma":[0.066642165,0.0013846625,0.69625115,0.0006867591,0.00010774782,0.0014527381,0.21681412,0.0108252205,0.0058353986],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99879104,0.0002029247,0.00023728328,0.0003541747,0.0003352958,0.00007928092],"domain_scores_gemma":[0.99553865,0.0018095409,0.00056764524,0.0009010235,0.0007453342,0.00043777341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037567788,0.0014368992,0.0012511188,0.0043064565,0.0013258215,0.0023779687,0.0021041606,0.0010454615,0.013000688],"category_scores_gemma":[0.009133165,0.0007340323,0.001435945,0.004279893,0.00070503564,0.0018350597,0.0024447225,0.0009688497,0.009238868],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042732283,0.00051886536,0.034634203,0.006583897,0.0007621548,0.002430073,0.0018073447,0.015319772,0.073056825,0.028411774,0.3579709,0.474231],"study_design_scores_gemma":[0.0009124899,0.00024718867,0.026773378,0.0010749538,0.0002762195,0.0022308063,0.0007812782,0.11431911,0.104229696,0.068328485,0.6804254,0.00040103344],"about_ca_topic_score_codex":0.0069573587,"about_ca_topic_score_gemma":0.008724658,"teacher_disagreement_score":0.013000688,"about_ca_system_score_codex":0.0011405379,"about_ca_system_score_gemma":0.0036781435,"threshold_uncertainty_score":0.04349166},"labels":[],"label_agreement":null},{"id":"W2030264436","doi":"10.2147/jmdh.s17564","title":"Clinical vocabulary as a boundary object in multidisciplinary care management of multiple chemical sensitivity, a complex and chronic condition","year":2011,"lang":"en","type":"article","venue":"Journal of Multidisciplinary Healthcare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Multidisciplinary approach; Terminology; Vocabulary; Medicine; Multiple chemical sensitivity; Health care; Object (grammar); Boundary (topology); Domain (mathematical analysis); Computer science; Artificial intelligence; Linguistics; Mathematics; Psychiatry","score_opus":0.04905071238811514,"score_gpt":0.3702257730136351,"score_spread":0.3211750606255199,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2030264436","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.75044155,0.0019751615,0.21843222,0.0041780947,0.00018134507,0.0024642108,0.0009597875,0.0004900507,0.020877652],"genre_scores_gemma":[0.8495803,0.00032807377,0.1473661,0.0002658121,0.000028061797,0.0010342876,0.0009182499,0.000045848457,0.00043327324],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.98638445,0.008962845,0.0020798114,0.0007996455,0.0015959125,0.00017733594],"domain_scores_gemma":[0.9412394,0.045085642,0.0065173367,0.0030177317,0.003455487,0.0006844318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015955796,0.00032719696,0.00050646695,0.0043179668,0.0012446844,0.003178692,0.0008874912,0.00085585465,0.0012914434],"category_scores_gemma":[0.052055236,0.0002242581,0.0006111539,0.0032930328,0.002704774,0.0068794917,0.003017364,0.0008080409,0.00019443402],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083823345,0.00054346875,0.22126563,0.0031492324,0.0002180354,0.001473952,0.12377137,0.0044111093,0.016600624,0.10555002,0.007949547,0.5142288],"study_design_scores_gemma":[0.0004955985,0.0020022336,0.3703176,0.0067931917,0.00086755893,0.008053795,0.17706393,0.080326825,0.026934868,0.19167924,0.1349644,0.0005006978],"about_ca_topic_score_codex":0.0023260952,"about_ca_topic_score_gemma":0.0019130564,"teacher_disagreement_score":0.015955796,"about_ca_system_score_codex":0.0023279663,"about_ca_system_score_gemma":0.004418842,"threshold_uncertainty_score":0.08438331},"labels":[],"label_agreement":null},{"id":"W2031815166","doi":"10.1016/j.jbi.2008.05.001","title":"yOWL: An ontology-driven knowledge base for yeast biologists","year":2008,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; Francis Crick Institute","keywords":"Computer science; Ontology; Knowledge base; Base (topology); Information retrieval; Data science; Artificial intelligence; Mathematics; Epistemology","score_opus":0.05117468690334019,"score_gpt":0.32675341496099475,"score_spread":0.27557872805765454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2031815166","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034076776,0.0012981086,0.7310621,0.0014793209,0.00036469393,0.0008671873,0.090674825,0.13253032,0.0076465984],"genre_scores_gemma":[0.12262571,0.0026769263,0.6496318,0.00093469064,0.00008902067,0.0009124841,0.21070084,0.0060134884,0.006415081],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99964094,0.00004182044,0.000061377694,0.00008156053,0.00014911985,0.000025193774],"domain_scores_gemma":[0.99828076,0.00078469305,0.00016643573,0.0003347473,0.00032491575,0.00010842032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011644163,0.0010760956,0.00087895145,0.0034224277,0.0007157326,0.0023028124,0.0021051643,0.0008751792,0.0051420988],"category_scores_gemma":[0.004966624,0.0007912657,0.0010086909,0.002487663,0.00038122197,0.0036851543,0.002556423,0.0014513396,0.0025304758],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015454553,0.0007767922,0.011504394,0.004236135,0.0007738375,0.0030194393,0.0011582638,0.04322144,0.042612158,0.032206498,0.1932257,0.66571987],"study_design_scores_gemma":[0.0005602061,0.0002700393,0.00818313,0.001140259,0.0008863744,0.001364866,0.00094302784,0.40239564,0.057849307,0.06852393,0.4575426,0.00034056697],"about_ca_topic_score_codex":0.008113201,"about_ca_topic_score_gemma":0.014504895,"teacher_disagreement_score":0.008113201,"about_ca_system_score_codex":0.00082760584,"about_ca_system_score_gemma":0.0020822326,"threshold_uncertainty_score":0.01720202},"labels":[],"label_agreement":null},{"id":"W2031950862","doi":"10.1093/nar/gkn296","title":"PolySearch: a web-based text mining system for extracting relationships between human diseases, genes, mutations, drugs and metabolites","year":2008,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":256,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Institute for Nanotechnology; University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; Genome Prairie; Genome Canada","keywords":"Identification (biology); Variety (cybernetics); Biology; Gene nomenclature; Rank (graph theory); Information retrieval; Computational biology; Computer science; Gene; Genomics; Genome; Bioinformatics; Genetics; Artificial intelligence; Nomenclature","score_opus":0.12613459816325415,"score_gpt":0.3758657254152483,"score_spread":0.24973112725199417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2031950862","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035218332,0.0028473653,0.23310989,0.0016675654,0.00029210956,0.001995701,0.3404619,0.37241453,0.011992602],"genre_scores_gemma":[0.06440181,0.0027568077,0.4247361,0.0008478189,0.0002622725,0.0019367915,0.47866204,0.009290775,0.017105503],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99908864,0.00014616123,0.00020605218,0.0002558764,0.00026105053,0.000042186344],"domain_scores_gemma":[0.99672323,0.0017839252,0.00051076984,0.00025709736,0.0004636026,0.00026137577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017754735,0.0024417893,0.0014701133,0.008784168,0.0009688838,0.0014640549,0.0011870255,0.0010442039,0.027068341],"category_scores_gemma":[0.005565825,0.00073973864,0.0015455221,0.004888177,0.0004874002,0.003887919,0.0017110525,0.0008454182,0.016586978],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023883414,0.0005267471,0.012057139,0.0070341947,0.0007542988,0.0024934665,0.0011372205,0.0035168105,0.052683726,0.0073674503,0.41501784,0.49502274],"study_design_scores_gemma":[0.0015614604,0.0015845017,0.04684801,0.0013706574,0.0010699957,0.00833856,0.0018864423,0.16219953,0.10926734,0.038170688,0.6270575,0.00064537616],"about_ca_topic_score_codex":0.0034654383,"about_ca_topic_score_gemma":0.0059003755,"teacher_disagreement_score":0.027068341,"about_ca_system_score_codex":0.0007359939,"about_ca_system_score_gemma":0.0022412434,"threshold_uncertainty_score":0.09055263},"labels":[],"label_agreement":null},{"id":"W2032039170","doi":"10.1038/npre.2011.6041.1","title":"Acknowledging contributions to online expert assistance","year":2011,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"","keywords":"Computer science; Data science","score_opus":0.01851932548750148,"score_gpt":0.3454918715416348,"score_spread":0.32697254605413334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2032039170","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.077643864,0.0053756847,0.086174086,0.1429999,0.021444589,0.001256026,0.0030749831,0.010299024,0.65173185],"genre_scores_gemma":[0.40771344,0.0030268428,0.083239384,0.02182352,0.010565375,0.0010629302,0.003359006,0.0035468566,0.4656626],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9612493,0.023849107,0.0017091015,0.0019155145,0.009186573,0.0020904664],"domain_scores_gemma":[0.7842873,0.12658678,0.012255206,0.019076971,0.035639014,0.02215475],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.025447946,0.00089143415,0.00055199803,0.0042813136,0.0037212789,0.008961,0.002283809,0.0045707542,0.11822957],"category_scores_gemma":[0.17148586,0.0004195313,0.0005733458,0.0025202767,0.0022369332,0.008215037,0.015326073,0.002167167,0.05226891],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027598208,0.00023374865,0.0039011964,0.0009946689,0.00003672767,0.0008390271,0.019733218,0.00039756927,0.0018559227,0.016346008,0.56588745,0.3894986],"study_design_scores_gemma":[0.00003446651,0.00011428244,0.0022153372,0.0006756694,0.000016886555,0.0004945915,0.010595466,0.0008158048,0.0010686605,0.011519652,0.9723866,0.000062688654],"about_ca_topic_score_codex":0.000614725,"about_ca_topic_score_gemma":0.0015287517,"teacher_disagreement_score":0.97455204,"about_ca_system_score_codex":0.0020350814,"about_ca_system_score_gemma":0.0037207275,"threshold_uncertainty_score":0.3955173},"labels":[],"label_agreement":null},{"id":"W2033776605","doi":"10.1177/1075547011401631","title":"Graphical and Computationally Intensive Techniques for Presenting and Disseminating Information About the Genetics of Disease","year":2011,"lang":"en","type":"article","venue":"Science Communication","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario College of Art and Design","funders":"","keywords":"Disease; Dissemination; Representation (politics); Genomics; Graphical display; Information Dissemination; Data science; Medical genetics; Computer science; Genetics; Biology; Genome; Medicine; World Wide Web; Political science; Gene","score_opus":0.028211697057606126,"score_gpt":0.3158339152394048,"score_spread":0.2876222181817987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2033776605","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009165918,0.0006257361,0.9896535,0.0015320466,0.000119531505,0.00012666674,0.00094696943,0.0027870724,0.0032918428],"genre_scores_gemma":[0.0427346,0.0037824453,0.9463706,0.0006130355,0.0002874667,0.0004454275,0.002413578,0.0004974835,0.0028553817],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943212,0.003140965,0.000521577,0.0006454365,0.0012218426,0.00014899134],"domain_scores_gemma":[0.97309506,0.018896481,0.0016130864,0.004792617,0.001277218,0.00032553606],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0059915762,0.0017661704,0.0011391899,0.00765379,0.0016330485,0.0073786452,0.002683491,0.0017599285,0.017028332],"category_scores_gemma":[0.03303902,0.00094255176,0.002382338,0.010321366,0.0024299014,0.010038778,0.004480485,0.0032388468,0.005132314],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022312207,0.00012946843,0.0013467935,0.0018299185,0.0002373134,0.00048100192,0.001658614,0.021846322,0.00602531,0.39157206,0.042807665,0.5318424],"study_design_scores_gemma":[0.000069065594,0.0000881079,0.0006980406,0.00046282628,0.00016638849,0.00097504933,0.0005723438,0.08627437,0.0065393364,0.78772116,0.11629282,0.0001405237],"about_ca_topic_score_codex":0.0030876396,"about_ca_topic_score_gemma":0.0042938697,"teacher_disagreement_score":0.99262136,"about_ca_system_score_codex":0.0015825469,"about_ca_system_score_gemma":0.0017562379,"threshold_uncertainty_score":0.05696547},"labels":[],"label_agreement":null},{"id":"W2034413846","doi":"10.1006/jbin.2001.1002","title":"Methods of Cognitive Analysis to Support the Design and Evaluation of Biomedical Systems: The Case of Clinical Practice Guidelines","year":2001,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":94,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"U.S. National Library of Medicine","keywords":"Computer science; Comprehension; Usability; Ambiguity; Cognition; Semantics (computer science); Automatic summarization; Knowledge management; Psychology; Artificial intelligence; Human–computer interaction","score_opus":0.3660530617101903,"score_gpt":0.5968159472395362,"score_spread":0.23076288552934593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2034413846","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050578685,0.17588095,0.78763086,0.009589772,0.00044203867,0.0007384418,0.00034088411,0.0010713197,0.019247841],"genre_scores_gemma":[0.050617334,0.06698793,0.8781709,0.0011237498,0.00020266492,0.0007249184,0.0003403861,0.00010136353,0.0017308526],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9828324,0.010758346,0.0011624205,0.0008631533,0.004160959,0.00022288441],"domain_scores_gemma":[0.9341868,0.056679055,0.0014038878,0.0025090596,0.004780135,0.0004410116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027917458,0.0021142985,0.0020902178,0.009893092,0.00092090614,0.007707363,0.0047725085,0.0028520243,0.0028271177],"category_scores_gemma":[0.042934187,0.0007881476,0.0013510919,0.00690158,0.005228462,0.0060662273,0.0027295114,0.0029537012,0.001067283],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000070383416,0.00015978933,0.0012941425,0.0067242235,0.00043746072,0.00021806791,0.0024758347,0.0050821714,0.00083387195,0.099964395,0.0065330034,0.87620664],"study_design_scores_gemma":[0.0002183775,0.00028161143,0.0060174237,0.013917073,0.0006742081,0.0018441909,0.0041117216,0.047293637,0.0049721356,0.6195246,0.30083305,0.00031198966],"about_ca_topic_score_codex":0.0079432605,"about_ca_topic_score_gemma":0.010015017,"teacher_disagreement_score":0.027917458,"about_ca_system_score_codex":0.0037983384,"about_ca_system_score_gemma":0.005758,"threshold_uncertainty_score":0.14764339},"labels":[],"label_agreement":null},{"id":"W2035595423","doi":"10.1093/nar/gkl037","title":"HubMed: a web-based biomedical literature search interface","year":2006,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network","funders":"University of Glasgow","keywords":"Metadata; Interface (matter); Information retrieval; World Wide Web; Citation; Computer science; Biology","score_opus":0.03006138763461771,"score_gpt":0.34970690947607075,"score_spread":0.31964552184145306,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035595423","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038420751,0.009636246,0.1449969,0.003115401,0.000591365,0.0032595112,0.5789964,0.20464642,0.050915685],"genre_scores_gemma":[0.016104452,0.0072912183,0.42344484,0.0033429242,0.00062550785,0.006328703,0.49469012,0.014002199,0.0341701],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987764,0.00036393225,0.00034090044,0.00014926722,0.00029948298,0.000070084876],"domain_scores_gemma":[0.99248224,0.0050950027,0.0005751695,0.0003013358,0.0010380734,0.0005081646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028986146,0.0020019065,0.0021934572,0.014682834,0.0006995735,0.0026943332,0.0025502404,0.0017912681,0.19289513],"category_scores_gemma":[0.013092405,0.0010311455,0.0012400664,0.0100901155,0.00034985444,0.0046216915,0.0044217,0.0013395968,0.08472303],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006656965,0.00011634642,0.00090647634,0.011770896,0.0002673171,0.0004965897,0.00038686043,0.00060772256,0.005803173,0.0062147365,0.781508,0.19125612],"study_design_scores_gemma":[0.00062294904,0.00010119241,0.0032213058,0.0012088644,0.00020953719,0.0007434887,0.00020440218,0.0032890688,0.0045566736,0.012600753,0.9730671,0.00017471645],"about_ca_topic_score_codex":0.0016546849,"about_ca_topic_score_gemma":0.0040671756,"teacher_disagreement_score":0.19289513,"about_ca_system_score_codex":0.00080080586,"about_ca_system_score_gemma":0.0022452562,"threshold_uncertainty_score":0.6452985},"labels":[],"label_agreement":null},{"id":"W2035715268","doi":"10.1038/nbt0805-925","title":"The Babel of genetic data terminology","year":2005,"lang":"en","type":"letter","venue":"Nature Biotechnology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":67,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"Terminology; Computational biology; Genetic data; Genetics; Biology; Evolutionary biology; Linguistics; Philosophy; Medicine","score_opus":0.015954835860823694,"score_gpt":0.2824158703567221,"score_spread":0.2664610344958984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035715268","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00030415299,0.0021469009,0.0038131021,0.98095095,0.008544451,0.00000824903,0.000069552785,0.00006449462,0.0040980345],"genre_scores_gemma":[0.012280553,0.0025435984,0.0067168837,0.9376727,0.031352304,0.00009543791,0.00010535339,0.00016203444,0.009071197],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98405993,0.008084534,0.0015366545,0.0015642493,0.003960081,0.00079453265],"domain_scores_gemma":[0.8532124,0.12296034,0.0036289298,0.005167789,0.011595534,0.0034349859],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.018275952,0.00062618224,0.0014760379,0.0023283283,0.0050078365,0.00892568,0.0023063142,0.0274724,0.0064523867],"category_scores_gemma":[0.097611286,0.0009853701,0.0011730442,0.0017698144,0.015704548,0.017712306,0.0052551655,0.046349257,0.004733468],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007764306,0.000010627654,0.00028070604,0.0000674711,0.000022999011,0.00065359316,0.00055434223,0.000115672156,0.00015930607,0.10833067,0.8689925,0.02073454],"study_design_scores_gemma":[0.00007716856,0.000023472474,0.00033534193,0.0004232206,0.000028380578,0.0015925938,0.00067960715,0.0012596941,0.00045118714,0.17698033,0.81807864,0.00007019399],"about_ca_topic_score_codex":0.006807119,"about_ca_topic_score_gemma":0.008964897,"teacher_disagreement_score":0.9910743,"about_ca_system_score_codex":0.0056456765,"about_ca_system_score_gemma":0.004802965,"threshold_uncertainty_score":0.09665358},"labels":[],"label_agreement":null},{"id":"W2035878947","doi":"10.1007/s10516-007-9021-0","title":"From the Universe of Knowledge to the Universe of Concepts: The Structural Revolution in Classification for Information Retrieval","year":2007,"lang":"en","type":"article","venue":"Axiomathes","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Ontology; Ideal (ethics); Universe; Library classification; Computer science; Order (exchange); Facet (psychology); Artificial intelligence; Epistemology; Library science; Philosophy; Physics; Psychology; Astronomy","score_opus":0.022630862375090578,"score_gpt":0.30237598387528986,"score_spread":0.27974512150019926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2035878947","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012176572,0.029251494,0.9155667,0.025670404,0.001302034,0.00016056254,0.0010517528,0.0006459156,0.014174576],"genre_scores_gemma":[0.21375594,0.021630935,0.74957263,0.0046686027,0.0030960322,0.0005810239,0.0020254783,0.000364422,0.004304932],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9925964,0.0032029331,0.000762208,0.0012257347,0.002034247,0.00017841821],"domain_scores_gemma":[0.9831993,0.0115855355,0.0008954867,0.0025610982,0.00129593,0.00046272165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008492568,0.0010166948,0.0019922773,0.0073681264,0.0029458588,0.010116694,0.0032424496,0.003162392,0.004326031],"category_scores_gemma":[0.027521426,0.0008629872,0.002346242,0.011018226,0.018285973,0.044605672,0.006927309,0.007973909,0.0013790026],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028141263,0.000025109975,0.00036320646,0.00029679688,0.000027247705,0.000045159308,0.0006937495,0.00044562353,0.000242352,0.949144,0.0032755877,0.04541299],"study_design_scores_gemma":[0.000009244297,0.000010371186,0.00014545256,0.000127402,0.000018129964,0.000101908765,0.00013902374,0.003331413,0.00016592062,0.97652155,0.019414214,0.00001544997],"about_ca_topic_score_codex":0.0024538927,"about_ca_topic_score_gemma":0.0014504236,"teacher_disagreement_score":0.010116694,"about_ca_system_score_codex":0.0027341023,"about_ca_system_score_gemma":0.0033029013,"threshold_uncertainty_score":0.04491359},"labels":[],"label_agreement":null},{"id":"W2037374844","doi":"10.3115/1572364.1572383","title":"Identifying interaction sentences from biological literature using automatically extracted patterns","year":2009,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Preprocessor; Sentence; Task (project management); Artificial intelligence; Identification (biology); Natural language processing; Pairwise comparison; Key (lock); Component (thermodynamics); Quality (philosophy); Pattern recognition (psychology)","score_opus":0.05341698554228467,"score_gpt":0.33701215180427563,"score_spread":0.28359516626199094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037374844","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.34972188,0.0063440893,0.57777125,0.0036484043,0.00065541547,0.0017096548,0.03724494,0.013148792,0.009755563],"genre_scores_gemma":[0.22241564,0.0016460756,0.7318346,0.00032925923,0.00040341512,0.0008482076,0.04039235,0.00031643707,0.0018139767],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99883324,0.00024545047,0.00025440697,0.00028330146,0.00032902716,0.00005459849],"domain_scores_gemma":[0.9903756,0.0055005485,0.00147073,0.000462465,0.0019728518,0.00021787979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015377757,0.0013059089,0.0009281806,0.0099957,0.00060592813,0.0013024948,0.0011201841,0.0011409764,0.0032189542],"category_scores_gemma":[0.010090706,0.00044288544,0.0010192338,0.0055359,0.00034322753,0.0027511362,0.0010727078,0.00071453006,0.0024460906],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008654081,0.0004931295,0.025531432,0.0068466435,0.00038865156,0.0045281094,0.0017617324,0.0038663105,0.17469724,0.0055994173,0.028389849,0.74703205],"study_design_scores_gemma":[0.0006033912,0.0020493986,0.16320753,0.0015916348,0.002091502,0.018497165,0.005217869,0.3790463,0.22570132,0.07375955,0.12770577,0.00052863255],"about_ca_topic_score_codex":0.000823083,"about_ca_topic_score_gemma":0.0016195771,"teacher_disagreement_score":0.0099957,"about_ca_system_score_codex":0.00052894856,"about_ca_system_score_gemma":0.0013540959,"threshold_uncertainty_score":0.010768473},"labels":[],"label_agreement":null},{"id":"W2037382435","doi":"10.1186/1471-213x-8-92","title":"An ontology for Xenopus anatomy and development","year":2008,"lang":"en","type":"article","venue":"BMC Developmental Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"National Institutes of Health; Eunice Kennedy Shriver National Institute of Child Health and Human Development; Alberta Heritage Foundation for Medical Research","keywords":"Xenopus; Biology; Ontology; Model organism; Computational biology; Controlled vocabulary; Bioinformatics; Computer science; Gene; Information retrieval; Genetics","score_opus":0.04007027699062416,"score_gpt":0.30840721166594276,"score_spread":0.2683369346753186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2037382435","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017424507,0.0050053503,0.83386445,0.00454839,0.0009547085,0.0013214075,0.08102889,0.010119709,0.045732606],"genre_scores_gemma":[0.093906984,0.008200326,0.7691609,0.0011640145,0.0003559339,0.0018861629,0.10814464,0.001272569,0.015908457],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99885035,0.0001794405,0.00028756994,0.00024657568,0.00033947075,0.00009661459],"domain_scores_gemma":[0.9985299,0.00054841884,0.0002502826,0.00020817557,0.00031771208,0.00014541023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001341326,0.0008032125,0.0005675117,0.0048087714,0.00177018,0.0021518501,0.0012008282,0.0012018741,0.0064003994],"category_scores_gemma":[0.0023545339,0.00047185237,0.002284937,0.0041495785,0.001084886,0.0045134593,0.0016345442,0.0018538666,0.001994541],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026102675,0.00017906247,0.0072126086,0.0034126143,0.00018322123,0.0015788943,0.0033325662,0.010892491,0.01967711,0.5798508,0.10272375,0.27069587],"study_design_scores_gemma":[0.000040544073,0.00004866505,0.0049424726,0.000642356,0.00014688469,0.0017947498,0.00055523036,0.0106159095,0.0034550922,0.07524235,0.9024381,0.00007761049],"about_ca_topic_score_codex":0.016912313,"about_ca_topic_score_gemma":0.018129088,"teacher_disagreement_score":0.016912313,"about_ca_system_score_codex":0.0029634316,"about_ca_system_score_gemma":0.005538568,"threshold_uncertainty_score":0.03362775},"labels":[],"label_agreement":null},{"id":"W2040520897","doi":"10.1109/titb.2011.2177097","title":"An Infobutton For Web 2.0 Clinical Discussions: The Knowledge Linkage Framework","year":2011,"lang":"en","type":"article","venue":"IEEE Transactions on Information Technology in Biomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Information retrieval; Lexicon; Linkage (software); Recall; Keyword search; Set (abstract data type); World Wide Web; Natural language processing; Psychology","score_opus":0.034026762428731226,"score_gpt":0.3437569239337854,"score_spread":0.30973016150505417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2040520897","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013039278,0.0012869047,0.9185432,0.0047823475,0.0002760183,0.002414148,0.012668667,0.037289724,0.009699575],"genre_scores_gemma":[0.030600537,0.00056236197,0.95374656,0.00038732227,0.000107627435,0.0010546881,0.009974309,0.0011004907,0.0024661142],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99147284,0.0034598254,0.0015015553,0.0010186638,0.0022273026,0.00031977118],"domain_scores_gemma":[0.9639239,0.024684992,0.0033189894,0.0028916276,0.0036252746,0.0015552882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017249566,0.0014003029,0.0018018942,0.02131679,0.0023749336,0.008631956,0.002246172,0.0024931072,0.0114844255],"category_scores_gemma":[0.056017596,0.0012510427,0.0019998518,0.01093115,0.001052147,0.012598644,0.006354857,0.0014656376,0.006810302],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009824949,0.0005255104,0.0069356402,0.0044858544,0.0003023788,0.00091853685,0.005505947,0.004143797,0.011095023,0.049273487,0.060351122,0.8554802],"study_design_scores_gemma":[0.00053780805,0.0008892657,0.008063514,0.004003694,0.0005156275,0.002372616,0.006305688,0.1416362,0.03173972,0.17751805,0.62579095,0.0006268203],"about_ca_topic_score_codex":0.0035002986,"about_ca_topic_score_gemma":0.004894608,"teacher_disagreement_score":0.02131679,"about_ca_system_score_codex":0.0020244515,"about_ca_system_score_gemma":0.0061140764,"threshold_uncertainty_score":0.091225505},"labels":[],"label_agreement":null},{"id":"W2042042646","doi":"10.1016/j.ipm.2010.03.010","title":"Mining and modeling linkage information from citation context for improving biomedical literature retrieval","year":2010,"lang":"en","type":"article","venue":"Information Processing & Management","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Ontario Ministry of Research and Innovation; Natural Sciences and Engineering Research Council of Canada; China Scholarship Council","keywords":"Computer science; Linkage (software); Information retrieval; Citation; Weighting; Graph; PageRank; Context (archaeology); Data mining; Probabilistic logic; Data science; Theoretical computer science; World Wide Web; Artificial intelligence","score_opus":0.010993876316661315,"score_gpt":0.25682058697189525,"score_spread":0.24582671065523393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2042042646","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22522053,0.0166072,0.7106661,0.0044513117,0.00080930704,0.0010460084,0.023138218,0.008757478,0.009303827],"genre_scores_gemma":[0.47960168,0.005778454,0.48468655,0.0003233228,0.00075081736,0.0009062094,0.024839891,0.0004236034,0.0026894894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9965829,0.00085158384,0.0006907898,0.0005582254,0.0011710948,0.0001452771],"domain_scores_gemma":[0.98665005,0.008308675,0.0013603839,0.00068146136,0.0027117922,0.00028772207],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0036266253,0.00090200437,0.0016502899,0.025757514,0.0018854485,0.003689358,0.0017714815,0.0017667641,0.002618508],"category_scores_gemma":[0.025534851,0.0004998856,0.0023616361,0.022947496,0.00043871882,0.0062827705,0.002009444,0.0012186901,0.0014052732],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009475365,0.0012244487,0.050463498,0.003910753,0.0012566347,0.0012792913,0.001342038,0.057818912,0.016959956,0.03498497,0.029575806,0.8002362],"study_design_scores_gemma":[0.00017863311,0.00034254964,0.015097158,0.00063095766,0.0020672546,0.0013005511,0.0008149025,0.8196497,0.017637635,0.1013946,0.040684644,0.00020147511],"about_ca_topic_score_codex":0.0075743734,"about_ca_topic_score_gemma":0.013963227,"teacher_disagreement_score":0.9742425,"about_ca_system_score_codex":0.0012538007,"about_ca_system_score_gemma":0.003871586,"threshold_uncertainty_score":0.019179702},"labels":[],"label_agreement":null},{"id":"W2043419849","doi":"10.3166/ria.18.111-137","title":"L'UMLS entre langue et ontologie : une approche pragmatique dans le domaine médical","year":2004,"lang":"fr","type":"article","venue":"Revue d intelligence artificielle","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Unified Medical Language System; Philosophy; Linguistics; Humanities; Computer science; Natural language processing","score_opus":0.027770793754921967,"score_gpt":0.2925595196639217,"score_spread":0.2647887259089997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2043419849","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021292483,0.0011478927,0.9849214,0.003915031,0.00022575945,0.00010950474,0.0002407227,0.0024710598,0.0048393514],"genre_scores_gemma":[0.04225901,0.0020260226,0.9463887,0.0013973603,0.00021404112,0.00022103632,0.00066624576,0.0007072178,0.006120407],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9923003,0.0036655946,0.0010804583,0.00095556746,0.0017838872,0.00021426508],"domain_scores_gemma":[0.99271506,0.0040329085,0.00052177,0.001311003,0.001150447,0.00026882635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011501988,0.0009879228,0.00092299806,0.0035096817,0.0018168576,0.0074204244,0.0017557197,0.0028496739,0.0039915163],"category_scores_gemma":[0.014623596,0.0010942568,0.0027659868,0.0027087426,0.0041901493,0.011193643,0.0045947554,0.0042650313,0.002088981],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013543201,0.00008233865,0.0013358403,0.00093147706,0.00012286114,0.00058585155,0.007619327,0.005590704,0.008726346,0.7001807,0.014655038,0.2600341],"study_design_scores_gemma":[0.000061660125,0.00007997552,0.00071105256,0.00082946464,0.00015389202,0.001210234,0.0012448794,0.042482037,0.008150956,0.34982008,0.5951303,0.0001254432],"about_ca_topic_score_codex":0.008312401,"about_ca_topic_score_gemma":0.007944053,"teacher_disagreement_score":0.011501988,"about_ca_system_score_codex":0.0023530559,"about_ca_system_score_gemma":0.004587438,"threshold_uncertainty_score":0.060829043},"labels":[],"label_agreement":null},{"id":"W2044438259","doi":"10.1109/bibm.2013.6732564","title":"Promoting electronic health record search through a time-aware approach","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"China Scholarship Council","keywords":"Relevance (law); Ranking (information retrieval); Electronic health record; Computer science; Similarity (geometry); Information retrieval; Centroid; Feature (linguistics); Health records; Interval (graph theory); Data mining; Artificial intelligence; Mathematics; Health care","score_opus":0.021250935061671243,"score_gpt":0.2901095212018692,"score_spread":0.268858586140198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044438259","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11706618,0.0043451274,0.8574435,0.0019752688,0.00020193945,0.0006506425,0.0006060949,0.00654238,0.011168787],"genre_scores_gemma":[0.57583123,0.0014008626,0.4140055,0.0005070531,0.00041710204,0.00027401163,0.0009790879,0.00031164841,0.006273577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99597543,0.0011889433,0.00039311012,0.0006729302,0.00153899,0.00023062504],"domain_scores_gemma":[0.9906466,0.0040809624,0.0014020974,0.0012722602,0.002206604,0.00039149486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025615566,0.00097372604,0.0012379833,0.00524267,0.001043387,0.0023982397,0.001841019,0.0013061421,0.0019596834],"category_scores_gemma":[0.015119881,0.00046179994,0.0008144179,0.0042559444,0.0004886751,0.005484463,0.001747469,0.0008384974,0.0015395135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011088102,0.0014278373,0.011290848,0.00069443305,0.000215048,0.0003345015,0.0009483115,0.03123956,0.07147099,0.010107036,0.009870905,0.8612917],"study_design_scores_gemma":[0.0003151904,0.0014428144,0.015091565,0.00011920582,0.0006699275,0.0019109218,0.001355071,0.86022544,0.04844866,0.02840368,0.041729365,0.00028815324],"about_ca_topic_score_codex":0.0046446095,"about_ca_topic_score_gemma":0.007173835,"teacher_disagreement_score":0.00524267,"about_ca_system_score_codex":0.00085569767,"about_ca_system_score_gemma":0.0019624506,"threshold_uncertainty_score":0.013546944},"labels":[],"label_agreement":null},{"id":"W2044500200","doi":"10.1145/1141753.1141878","title":"Unsupervised structure discovery for biodiversity information","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Biodiversity; Computer science; Data science; Information retrieval; Ecology; Biology","score_opus":0.0061073602124201865,"score_gpt":0.20816042806575175,"score_spread":0.20205306785333155,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044500200","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08970571,0.0027830217,0.8764328,0.0011983487,0.00014666752,0.0005039319,0.01424244,0.005843932,0.0091432],"genre_scores_gemma":[0.35786295,0.0010334856,0.5990677,0.00022063963,0.00014057475,0.00050658605,0.03570189,0.000433232,0.005033021],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99892503,0.00029484674,0.00006676917,0.00027950262,0.00034119465,0.00009258425],"domain_scores_gemma":[0.99722916,0.0014223114,0.00026216314,0.0005302985,0.00042373926,0.00013241006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010222027,0.0006166065,0.0010475203,0.0051321243,0.0009572447,0.001267763,0.0012856942,0.0007951546,0.004807131],"category_scores_gemma":[0.004722303,0.0003674361,0.0011552821,0.0038294918,0.0006163266,0.0021426396,0.0015200567,0.0008694379,0.0017349123],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006765942,0.00048974314,0.012526919,0.0011246311,0.0004832206,0.00038083334,0.00026256396,0.02976298,0.031010212,0.023353908,0.0344688,0.8654595],"study_design_scores_gemma":[0.00017077346,0.0002666321,0.011458635,0.00022322946,0.00035609392,0.00082118675,0.00027447284,0.7706458,0.025614237,0.15761851,0.032461684,0.00008877007],"about_ca_topic_score_codex":0.0025432524,"about_ca_topic_score_gemma":0.0052912706,"teacher_disagreement_score":0.0051321243,"about_ca_system_score_codex":0.00077709684,"about_ca_system_score_gemma":0.0014674161,"threshold_uncertainty_score":0.016081512},"labels":[],"label_agreement":null},{"id":"W2044833006","doi":"10.1186/1471-2105-12-s3-s1","title":"Building a biomedical tokenizer using the token lattice design pattern and the adapted Viterbi algorithm","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Lexical analysis; Computer science; Security token; Viterbi algorithm; Artificial intelligence; Domain (mathematical analysis); Classifier (UML); Natural language processing; Hidden Markov model","score_opus":0.060771486185793054,"score_gpt":0.2754870983270212,"score_spread":0.2147156121412281,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2044833006","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00686898,0.00007498647,0.9856205,0.00017852247,0.00007670069,0.00015342208,0.00018474003,0.0059080594,0.00093417807],"genre_scores_gemma":[0.0477733,0.000057906156,0.9484025,0.00014890626,0.000028059594,0.0002295472,0.0006339634,0.000627663,0.0020982162],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9972269,0.00079477753,0.00039708894,0.00075274333,0.00064591673,0.00018262354],"domain_scores_gemma":[0.99596304,0.0017759674,0.0003604054,0.00063787017,0.0011071416,0.00015557735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003681378,0.0008568852,0.0010245513,0.0016570196,0.0009850726,0.0016866559,0.0025501312,0.0015303682,0.0069899303],"category_scores_gemma":[0.008888351,0.00074382656,0.00089253864,0.0017507826,0.0015229877,0.002720439,0.0017924603,0.0019719547,0.006255101],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012598726,0.00025608827,0.0041310987,0.00073083886,0.00012852561,0.00050842226,0.00064228073,0.065716594,0.05401304,0.049267992,0.018462995,0.8048822],"study_design_scores_gemma":[0.00021244201,0.00028620014,0.00046762364,0.000067395726,0.00008615672,0.0005892296,0.00015106043,0.80758506,0.122158475,0.037456412,0.030855313,0.00008475347],"about_ca_topic_score_codex":0.0029874314,"about_ca_topic_score_gemma":0.004175759,"teacher_disagreement_score":0.0069899303,"about_ca_system_score_codex":0.0013148217,"about_ca_system_score_gemma":0.004042481,"threshold_uncertainty_score":0.023383617},"labels":[],"label_agreement":null},{"id":"W2045484586","doi":"10.1016/j.ymeth.2014.11.022","title":"Assessment of curated phenotype mining in neuropsychiatric disorder literature","year":2014,"lang":"en","type":"article","venue":"Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; University of Ottawa","funders":"Bundesministerium für Bildung und Forschung","keywords":"Phenotype; Disease; Clinical phenotype; Gene; Bioinformatics; Medicine; Computational biology; Biology; Genetics; Pathology","score_opus":0.013863338318895467,"score_gpt":0.3806454225225763,"score_spread":0.36678208420368086,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2045484586","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60807705,0.022794107,0.20676017,0.0044545345,0.00065804407,0.0022427652,0.12558483,0.012742759,0.016685808],"genre_scores_gemma":[0.5952471,0.004168081,0.27029842,0.0008463187,0.0001751595,0.0013249206,0.12441474,0.00087467645,0.0026506966],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98731333,0.0039578304,0.002930172,0.0017775502,0.0036634312,0.000357671],"domain_scores_gemma":[0.9299252,0.047642037,0.0043531978,0.0062272144,0.010814642,0.0010376661],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013721946,0.0008862545,0.0011529531,0.018958613,0.0013049907,0.0039015294,0.0019614366,0.0011711495,0.0029610125],"category_scores_gemma":[0.06850995,0.00030484542,0.0018171231,0.008621284,0.00051193795,0.00218904,0.0030503771,0.00071999425,0.0010255284],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020866843,0.0010061981,0.2631438,0.019014128,0.0044331555,0.0052332943,0.0032901636,0.013172567,0.02594465,0.009566977,0.03486268,0.61824566],"study_design_scores_gemma":[0.00087542576,0.0011401009,0.38741684,0.009322949,0.010993869,0.015982918,0.0054410654,0.22196041,0.06962455,0.036736406,0.24014753,0.00035793686],"about_ca_topic_score_codex":0.0030541338,"about_ca_topic_score_gemma":0.0073544374,"teacher_disagreement_score":0.98627806,"about_ca_system_score_codex":0.0010167863,"about_ca_system_score_gemma":0.00456161,"threshold_uncertainty_score":0.07256943},"labels":[],"label_agreement":null},{"id":"W2046381118","doi":"10.1046/j.1528-1157.2001.22001.x","title":"Glossary of Descriptive Terminology for Ictal Semiology: Report of the ILAE Task Force on Classification and Terminology","year":2001,"lang":"en","type":"article","venue":"Epilepsia","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":859,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"London Health Sciences Centre","funders":"","keywords":"Semiology; Psychology; Automatism (medicine); Ictal; Epilepsy; Tonic (physiology); Audiology; Neuroscience; Medicine","score_opus":0.03703950545523505,"score_gpt":0.29130683590673884,"score_spread":0.2542673304515038,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2046381118","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002681316,0.06440557,0.08306812,0.020757755,0.018513292,0.0069030025,0.4995331,0.004320254,0.29981762],"genre_scores_gemma":[0.019630786,0.09258265,0.12349549,0.019067101,0.011434427,0.020094953,0.60565275,0.006415304,0.1016266],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9921807,0.0020529109,0.002413511,0.00061075215,0.0024523751,0.00028969382],"domain_scores_gemma":[0.9620165,0.017819496,0.0035980148,0.0027082926,0.01318432,0.00067337696],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073687546,0.0016594582,0.0020972856,0.017302956,0.001862861,0.006518082,0.0031337587,0.0015798256,0.10401287],"category_scores_gemma":[0.04397553,0.00087008177,0.0012317119,0.023346165,0.0016903549,0.0060488894,0.0029376056,0.00332704,0.09007732],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035907924,0.000016953843,0.00031463854,0.0022923504,0.000008046706,0.000073057374,0.00039771968,0.0001503484,0.00037753867,0.013134291,0.9446812,0.038517814],"study_design_scores_gemma":[0.0000072302564,0.00000879889,0.000818513,0.0023356546,0.000008743082,0.00016312471,0.00024807904,0.00013262982,0.00012260511,0.003974378,0.9921588,0.00002157206],"about_ca_topic_score_codex":0.010323957,"about_ca_topic_score_gemma":0.0072749546,"teacher_disagreement_score":0.10401287,"about_ca_system_score_codex":0.004302153,"about_ca_system_score_gemma":0.006819006,"threshold_uncertainty_score":0.34795773},"labels":[],"label_agreement":null},{"id":"W2047221535","doi":"10.1016/j.jclinepi.2011.10.014","title":"Search filters can find some but not all knowledge translation articles in MEDLINE: an analytic survey","year":2012,"lang":"en","type":"review","venue":"Journal of Clinical Epidemiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":41,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; University of Toronto; McMaster University","funders":"Canadian Institutes of Health Research","keywords":"Terminology; MEDLINE; Filter (signal processing); Information retrieval; Computer science; Field (mathematics); Sensitivity (control systems); Translation (biology); Medical physics; Medicine; Mathematics; Linguistics","score_opus":0.727354163457325,"score_gpt":0.5872265397795597,"score_spread":0.1401276236777652,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2047221535","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056494225,0.9028265,0.0035892127,0.005090807,0.00014841983,0.00044534763,0.025049755,0.00022497089,0.006130653],"genre_scores_gemma":[0.19367012,0.7504391,0.01744226,0.004154527,0.00032433233,0.0007480856,0.030845957,0.00018937254,0.0021862774],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98939425,0.0023535483,0.0042635123,0.0008287184,0.002813144,0.00034686958],"domain_scores_gemma":[0.7953234,0.17154819,0.01673732,0.0024469695,0.01293741,0.0010066928],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015187743,0.0010118199,0.0028156783,0.04377338,0.00082366075,0.0028959657,0.001328777,0.0014302159,0.004660212],"category_scores_gemma":[0.098382644,0.00051879353,0.0026463217,0.049150553,0.00066293933,0.004603708,0.0015513285,0.0006240413,0.0013909871],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015021618,0.00016701818,0.06604442,0.1616299,0.007402988,0.0011572934,0.0030757557,0.00039563433,0.0031478724,0.00251267,0.03730013,0.71566415],"study_design_scores_gemma":[0.00065202755,0.00082857464,0.29556188,0.19446306,0.068765365,0.010729781,0.006143783,0.0013794724,0.0064549437,0.008674508,0.40607646,0.0002702158],"about_ca_topic_score_codex":0.0038345687,"about_ca_topic_score_gemma":0.014266805,"teacher_disagreement_score":0.98481226,"about_ca_system_score_codex":0.0012883327,"about_ca_system_score_gemma":0.0050770855,"threshold_uncertainty_score":0.08032143},"labels":[],"label_agreement":null},{"id":"W2048951946","doi":"10.7472/jksii.2014.15.4.21","title":"Evaluation of Usefulness of the Protein Drug Feature Information Filed","year":2014,"lang":"en","type":"article","venue":"Journal of Internet Computing and services","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"DrugBank; Usability; Service (business); Computer science; Data science; World Wide Web; Drug; Medicine; Business","score_opus":0.009911378941748664,"score_gpt":0.25168716822443027,"score_spread":0.24177578928268162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2048951946","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.891416,0.0034458116,0.011860351,0.0014721238,0.0007503443,0.0007710542,0.052345634,0.020401103,0.0175376],"genre_scores_gemma":[0.92269725,0.00097187486,0.021287516,0.00019835465,0.00017938111,0.0003117853,0.04745712,0.00079777953,0.006098859],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9951049,0.00084079837,0.00075606786,0.000918449,0.0021056717,0.00027400337],"domain_scores_gemma":[0.9314898,0.044712186,0.003443427,0.004211988,0.014362905,0.0017797215],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061212317,0.00087312184,0.0007071438,0.010428824,0.00075770763,0.0020978344,0.0011506189,0.0011571451,0.0043660537],"category_scores_gemma":[0.050237786,0.00021553022,0.00072328805,0.005348999,0.00052152644,0.0038485972,0.001190415,0.00056943414,0.0021453581],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006581762,0.0014864813,0.32633364,0.0024936644,0.0004328172,0.0016516275,0.0017854079,0.012367595,0.009663234,0.0019073713,0.07401796,0.5612785],"study_design_scores_gemma":[0.0007544298,0.0037164662,0.6056054,0.000763606,0.0010080638,0.0030945193,0.0053692358,0.23285894,0.049334753,0.0031408714,0.09382434,0.00052922854],"about_ca_topic_score_codex":0.008729789,"about_ca_topic_score_gemma":0.004745202,"teacher_disagreement_score":0.010428824,"about_ca_system_score_codex":0.001071397,"about_ca_system_score_gemma":0.0012193164,"threshold_uncertainty_score":0.032372534},"labels":[],"label_agreement":null},{"id":"W2049645944","doi":"10.1016/j.jbi.2012.02.012","title":"Lexical patterns, features and knowledge resources for coreference resolution in clinical notes","year":2012,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"U.S. National Library of Medicine","keywords":"Coreference; Computer science; Natural language processing; Artificial intelligence; Baseline (sea); Measure (data warehouse); Resolution (logic); Variety (cybernetics); Recall; Precision and recall; Information retrieval; Data mining; Linguistics","score_opus":0.04965589118374901,"score_gpt":0.3658778250269115,"score_spread":0.3162219338431625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2049645944","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5658287,0.0045145764,0.36401483,0.0042502396,0.00029722656,0.0013993765,0.03666746,0.0063252663,0.016702399],"genre_scores_gemma":[0.7340262,0.0006855921,0.24100007,0.00030031387,0.00007462027,0.00055176055,0.021288779,0.00029693716,0.001775725],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99769324,0.0006837203,0.00063970167,0.00034812614,0.00048065695,0.00015448533],"domain_scores_gemma":[0.99036366,0.0072140205,0.00059856253,0.00066214224,0.0009274795,0.00023411059],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024287237,0.00039726234,0.0007129411,0.0071811033,0.0013110446,0.0024905659,0.0010546112,0.0010264233,0.0047121616],"category_scores_gemma":[0.014516821,0.00034416546,0.0008602607,0.004535828,0.00061361544,0.0036958645,0.0020993904,0.00094790943,0.0011330016],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023336727,0.0009112032,0.06531604,0.0024761988,0.00045320243,0.0043496937,0.0034918983,0.0072419504,0.03798751,0.031090282,0.029369581,0.8149787],"study_design_scores_gemma":[0.0008619178,0.00097623107,0.13873133,0.0032081504,0.003183502,0.016749732,0.01641947,0.39384237,0.10292186,0.1871049,0.13550135,0.0004992119],"about_ca_topic_score_codex":0.0034963125,"about_ca_topic_score_gemma":0.006019255,"teacher_disagreement_score":0.0071811033,"about_ca_system_score_codex":0.0007981414,"about_ca_system_score_gemma":0.0019886058,"threshold_uncertainty_score":0.01576376},"labels":[],"label_agreement":null},{"id":"W2050263604","doi":"10.1186/1471-2105-12-435","title":"PESCADOR, a web-based tool to assist text-mining of biointeractions extracted from PubMed queries","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung; Ottawa Hospital Research Institute","keywords":"Computer science; Relevance (law); Function (biology); Set (abstract data type); Resource (disambiguation); Mechanism (biology); Biological network; Web resource; Focus (optics); Information retrieval; Interaction network; Computational biology; World Wide Web; Bioinformatics; Biology; Gene","score_opus":0.03956506832173659,"score_gpt":0.25544508606942734,"score_spread":0.21588001774769075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050263604","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031851668,0.006715608,0.13701417,0.0016282506,0.00026317578,0.0025279915,0.59683174,0.20804729,0.015120069],"genre_scores_gemma":[0.06151467,0.0054515977,0.48027468,0.00071631,0.00015332086,0.003014268,0.43694007,0.006485788,0.005449313],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988814,0.00017748424,0.00034153316,0.00026380815,0.0002866865,0.000049111142],"domain_scores_gemma":[0.9927685,0.004818943,0.0010264898,0.00040826973,0.00071941415,0.0002583142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027601584,0.0025831058,0.001502091,0.025494453,0.000977749,0.0024585861,0.0015103718,0.0009965108,0.026666094],"category_scores_gemma":[0.0138520235,0.0006853131,0.0014986115,0.014271244,0.0004886928,0.0037454332,0.002474182,0.0009290737,0.0075427564],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023106877,0.00038351293,0.013337559,0.044333026,0.001619598,0.005275517,0.003482505,0.004629803,0.043283127,0.014948886,0.34832916,0.5180666],"study_design_scores_gemma":[0.0010646009,0.0004841911,0.029880453,0.003468064,0.0013179942,0.00620949,0.0017125317,0.03884552,0.030408662,0.020427028,0.86575234,0.0004291598],"about_ca_topic_score_codex":0.0039185015,"about_ca_topic_score_gemma":0.008864876,"teacher_disagreement_score":0.026666094,"about_ca_system_score_codex":0.0012181399,"about_ca_system_score_gemma":0.0031009582,"threshold_uncertainty_score":0.08920699},"labels":[],"label_agreement":null},{"id":"W2050554833","doi":"10.1038/nbt.1411","title":"Promoting coherent minimum reporting guidelines for biological and biomedical investigations: the MIBBI project","year":2008,"lang":"en","type":"article","venue":"Nature Biotechnology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":562,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"Natural Environment Research Council; National Institute of Biomedical Imaging and Bioengineering; National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council","keywords":"Extant taxon; Resource (disambiguation); Knowledge management; Business; Management science; Data science; Computer science; Biology; Engineering; Evolutionary biology","score_opus":0.1045177839257098,"score_gpt":0.37075066261321704,"score_spread":0.2662328786875072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050554833","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0106343245,0.005154057,0.8490278,0.09627615,0.0036041306,0.015193135,0.0034312292,0.004672381,0.012006839],"genre_scores_gemma":[0.015375169,0.0016505627,0.9602739,0.004388783,0.00036213716,0.009946713,0.0057405294,0.0006243709,0.0016378936],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.46924016,0.3314427,0.098536596,0.010977409,0.08354407,0.0062591573],"domain_scores_gemma":[0.15947817,0.36567634,0.078576684,0.12262851,0.2514776,0.022162726],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.61506563,0.0022484611,0.0031454007,0.024826236,0.008678718,0.019070351,0.015292539,0.010189712,0.0025809878],"category_scores_gemma":[0.66585004,0.0034615889,0.0032090398,0.021360578,0.007691522,0.022082962,0.024275385,0.016567342,0.0021804587],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006834367,0.0010812493,0.009838993,0.010576865,0.00042398417,0.0006398832,0.016773984,0.0051634875,0.007811187,0.23491152,0.2566675,0.45542797],"study_design_scores_gemma":[0.0006707129,0.000887837,0.011143897,0.021762593,0.00061394833,0.0008846633,0.008400812,0.019830734,0.01824082,0.1970438,0.71974677,0.000773487],"about_ca_topic_score_codex":0.013263113,"about_ca_topic_score_gemma":0.008454461,"teacher_disagreement_score":0.38493437,"about_ca_system_score_codex":0.01508711,"about_ca_system_score_gemma":0.1228542,"threshold_uncertainty_score":0.47469264},"labels":[],"label_agreement":null},{"id":"W2050567956","doi":"10.3115/1654415.1654432","title":"Postnominal prepositional phrase attachment in proteomics","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Heuristics; Natural language processing; Noun phrase; Artificial intelligence; Nominalization; Phrase; Set (abstract data type); Test set; Information retrieval; Programming language; Noun","score_opus":0.005977868209777394,"score_gpt":0.25102877536951435,"score_spread":0.24505090715973696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2050567956","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17864259,0.0037473592,0.77950966,0.0012222342,0.00061643845,0.00069661636,0.0052445615,0.015398578,0.014922038],"genre_scores_gemma":[0.5048911,0.0014088298,0.4779651,0.00048579974,0.00025634613,0.00024851534,0.008297448,0.0009985708,0.005448228],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978186,0.00077256653,0.0003076999,0.0005491915,0.00042775422,0.00012413763],"domain_scores_gemma":[0.99026066,0.0064630765,0.0010119815,0.0009266537,0.0011109145,0.0002266426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029887268,0.0011501438,0.00078949874,0.0030561266,0.0014327958,0.0024791723,0.0011841852,0.0011927567,0.005757216],"category_scores_gemma":[0.0095309755,0.00076251256,0.00082378276,0.0027785536,0.0010953717,0.0035045291,0.0014402145,0.0013688633,0.004552203],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016913869,0.00034201218,0.023588467,0.003914632,0.00018703028,0.002669711,0.0059919455,0.0071002105,0.16498011,0.04111386,0.035060946,0.7133598],"study_design_scores_gemma":[0.00026990252,0.00079451804,0.065584704,0.0012364719,0.00056092313,0.007927815,0.007061941,0.22226653,0.2968627,0.1339928,0.26295015,0.00049157214],"about_ca_topic_score_codex":0.0017318513,"about_ca_topic_score_gemma":0.0027782538,"teacher_disagreement_score":0.005757216,"about_ca_system_score_codex":0.0008364246,"about_ca_system_score_gemma":0.0014280853,"threshold_uncertainty_score":0.01925981},"labels":[],"label_agreement":null},{"id":"W2051581953","doi":"10.1145/2064696.2064705","title":"Semantic text mining for lignocellulose research","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Génome Québec; Genome Canada","keywords":"Computer science; Semantic Web; Identification (biology); Social Semantic Web; World Wide Web; Semantic analytics; Ontology; Semantic Web Stack; Interface (matter); Semantics (computer science); Data science; Information retrieval","score_opus":0.17773506746946643,"score_gpt":0.36992092372058644,"score_spread":0.19218585625112,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051581953","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026086178,0.007636612,0.9234854,0.004981354,0.0003229548,0.0013376903,0.016490223,0.009478343,0.01018128],"genre_scores_gemma":[0.0659817,0.004224849,0.9030245,0.0003522247,0.00014402543,0.0006339936,0.023868427,0.00028301674,0.0014872644],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974552,0.0009097262,0.0006081863,0.00032823838,0.0006331893,0.000065463806],"domain_scores_gemma":[0.99507034,0.0026191778,0.0005846004,0.0005944961,0.0009883745,0.00014309156],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004034079,0.0007533894,0.00077120063,0.011612019,0.0016113996,0.002614281,0.0010703471,0.0009174769,0.003395552],"category_scores_gemma":[0.00950956,0.00035426204,0.0014247375,0.009477683,0.0010300173,0.0045841685,0.0016243499,0.0009350539,0.0015191488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027903845,0.00042473443,0.0044838935,0.0055108666,0.00026987444,0.0014206519,0.001253125,0.013679082,0.030772587,0.11600692,0.03400589,0.7918934],"study_design_scores_gemma":[0.00014304541,0.00014927304,0.0065016896,0.0015538268,0.00037260383,0.001595192,0.0018736939,0.2016679,0.055832475,0.39830017,0.33186078,0.00014933142],"about_ca_topic_score_codex":0.002244702,"about_ca_topic_score_gemma":0.0033262037,"teacher_disagreement_score":0.011612019,"about_ca_system_score_codex":0.0015164793,"about_ca_system_score_gemma":0.0039400007,"threshold_uncertainty_score":0.02133447},"labels":[],"label_agreement":null},{"id":"W2051702682","doi":"10.1080/02701960.2011.598973","title":"A Landscape for Training in Dementia Knowledge Translation (DKT)","year":2011,"lang":"en","type":"article","venue":"Gerontology & Geriatrics Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of British Columbia Hospital; NeuroDevNet; University of British Columbia","funders":"Canadian Institutes of Health Research","keywords":"Knowledge translation; Dementia; Economic shortage; Knowledge transfer; Bridge (graph theory); Knowledge management; Training (meteorology); Quality (philosophy); Knowledge sharing; Medical education; Medicine; Psychology; Nursing; Computer science; Geography","score_opus":0.12750176586374143,"score_gpt":0.33826829635521644,"score_spread":0.210766530491475,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051702682","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062861047,0.014935087,0.049545124,0.8996851,0.00213501,0.0007510524,0.00022061766,0.0008710907,0.025570763],"genre_scores_gemma":[0.24070038,0.03334736,0.55327886,0.14735754,0.004731122,0.006094157,0.001414182,0.00087152154,0.012204925],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.82438046,0.12946394,0.012608653,0.010561224,0.01316836,0.009817302],"domain_scores_gemma":[0.51879925,0.32017097,0.017462105,0.030276451,0.053605355,0.059685893],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.262492,0.0014699586,0.0022986946,0.008346696,0.011001046,0.029681055,0.010447269,0.024536395,0.027270935],"category_scores_gemma":[0.2295246,0.0019086973,0.0026991782,0.00711449,0.019891378,0.043254536,0.03726887,0.020098649,0.010588402],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038675033,0.0018159127,0.0055188974,0.0065352474,0.00011570488,0.00090282183,0.016676776,0.0018934999,0.0029079856,0.2309299,0.122165784,0.6101507],"study_design_scores_gemma":[0.00041593757,0.0008063352,0.008368918,0.01387965,0.00008797876,0.0015887747,0.039205536,0.005322366,0.0014403124,0.43937358,0.48916087,0.00034978887],"about_ca_topic_score_codex":0.010590281,"about_ca_topic_score_gemma":0.0056357454,"teacher_disagreement_score":0.262492,"about_ca_system_score_codex":0.014355802,"about_ca_system_score_gemma":0.10679608,"threshold_uncertainty_score":0.90947866},"labels":[],"label_agreement":null},{"id":"W2051847847","doi":"10.1016/j.ijmedinf.2011.02.006","title":"Implications of SNOMED CT versioning","year":2011,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Systematized Nomenclature of Medicine; Terminology; Documentation; Computer science; Hierarchy; Information retrieval; Component (thermodynamics); Medicine; Programming language; Linguistics","score_opus":0.03779408847711671,"score_gpt":0.32307488382831484,"score_spread":0.2852807953511981,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051847847","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33606473,0.011325806,0.40060404,0.12702812,0.004529919,0.00094861165,0.017645346,0.0039276006,0.09792582],"genre_scores_gemma":[0.77171963,0.0030929882,0.20441294,0.006562015,0.0010890178,0.00019569234,0.0070098867,0.001361092,0.0045567052],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9649993,0.018540658,0.0033802919,0.0028275384,0.009389535,0.0008627166],"domain_scores_gemma":[0.7074476,0.2156151,0.011106177,0.032856595,0.030723583,0.0022509666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031114623,0.0009485124,0.0008700549,0.006514921,0.002357991,0.009923981,0.004105862,0.003063261,0.007930262],"category_scores_gemma":[0.20909719,0.00073242164,0.0014420692,0.009104835,0.0045247125,0.015336439,0.0032084556,0.0039128643,0.0013474181],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00243238,0.00050033885,0.06972437,0.0016676006,0.00066815503,0.0058697723,0.008507772,0.023213508,0.00822691,0.46257046,0.037457805,0.37916097],"study_design_scores_gemma":[0.00021567411,0.00021752741,0.017090397,0.0015825765,0.0010735545,0.0077831806,0.0050871186,0.05674149,0.009249456,0.8006128,0.1000826,0.00026374287],"about_ca_topic_score_codex":0.012207161,"about_ca_topic_score_gemma":0.008542097,"teacher_disagreement_score":0.031114623,"about_ca_system_score_codex":0.0023525613,"about_ca_system_score_gemma":0.005737347,"threshold_uncertainty_score":0.1645518},"labels":[],"label_agreement":null},{"id":"W2051873544","doi":"10.1016/j.jbi.2012.11.007","title":"State of the art and open challenges in community-driven knowledge curation","year":2012,"lang":"en","type":"letter","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Data curation; Computer science; State (computer science); World Wide Web; Data science; Knowledge management","score_opus":0.06846499735600492,"score_gpt":0.3251504983102258,"score_spread":0.2566855009542209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2051873544","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008413218,0.012167945,0.014062702,0.96603626,0.004593774,0.000035435627,0.00015494085,0.00013793018,0.0019696604],"genre_scores_gemma":[0.077293694,0.059773248,0.13948435,0.62200874,0.09218802,0.000684433,0.001776036,0.0004614142,0.0063300524],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.946033,0.02605659,0.0051702647,0.0036840471,0.016990252,0.0020659207],"domain_scores_gemma":[0.5825625,0.34359932,0.008066151,0.015060997,0.041902572,0.008808384],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07821404,0.0005416426,0.0018643846,0.0027545479,0.0043823826,0.014223789,0.0063599567,0.018054524,0.0073801544],"category_scores_gemma":[0.17752127,0.0007998892,0.0012292321,0.0033131987,0.01119217,0.03180767,0.008938923,0.017071454,0.0039732037],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033045947,0.00019163193,0.0022066229,0.0027454665,0.00015712407,0.000798533,0.0021670896,0.0011203345,0.0009827854,0.09681258,0.5197537,0.3727337],"study_design_scores_gemma":[0.00009664201,0.00008502122,0.0009740937,0.002389447,0.000060351194,0.0014052651,0.0028421103,0.009920631,0.00075180904,0.3099704,0.6713474,0.00015674405],"about_ca_topic_score_codex":0.0049565365,"about_ca_topic_score_gemma":0.010050086,"teacher_disagreement_score":0.92178595,"about_ca_system_score_codex":0.0046259533,"about_ca_system_score_gemma":0.0091927685,"threshold_uncertainty_score":0.41364032},"labels":[],"label_agreement":null},{"id":"W2052057394","doi":"10.1186/1471-2105-10-313","title":"Social tagging in the life sciences: characterizing a new metadata resource for bioinformatics","year":2009,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; St. Paul's Hospital","funders":"Natural Sciences and Engineering Research Council of Canada; University of British Columbia","keywords":"Metadata; Computer science; World Wide Web; Information retrieval; Search engine indexing; Resource (disambiguation); Metadata repository; Annotation; Data science; Artificial intelligence","score_opus":0.07446472868360497,"score_gpt":0.32471233606516253,"score_spread":0.25024760738155755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052057394","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70475197,0.009307322,0.22360761,0.013401619,0.00041727594,0.0010080183,0.009739191,0.0014435044,0.036323458],"genre_scores_gemma":[0.84188306,0.0023995892,0.14588368,0.00080593856,0.00040470195,0.0006526232,0.00572562,0.00021493642,0.002029745],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98948765,0.0048763934,0.0013958694,0.0012225463,0.0025637795,0.00045380456],"domain_scores_gemma":[0.88303125,0.06537225,0.01872083,0.015511312,0.013251343,0.0041129473],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.017855853,0.00041081195,0.00069851923,0.015227906,0.0033695623,0.007909286,0.0013534158,0.0013971466,0.0018513262],"category_scores_gemma":[0.058259685,0.00030276782,0.000994349,0.02155015,0.0034699256,0.008719558,0.0051990543,0.0010335727,0.0009160915],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007758534,0.00034874637,0.4274249,0.003520155,0.0004625332,0.0007162201,0.02445316,0.0039918837,0.02158765,0.09505815,0.0128805395,0.40878022],"study_design_scores_gemma":[0.00011171478,0.0007960201,0.3668656,0.0025844176,0.00078529347,0.0034752022,0.03283023,0.042651188,0.03631346,0.25476158,0.25818235,0.0006430397],"about_ca_topic_score_codex":0.0040071346,"about_ca_topic_score_gemma":0.0075864797,"teacher_disagreement_score":0.9920907,"about_ca_system_score_codex":0.0036621387,"about_ca_system_score_gemma":0.004967451,"threshold_uncertainty_score":0.09443188},"labels":[],"label_agreement":null},{"id":"W2052547205","doi":"10.1016/j.ijmedinf.2005.06.010","title":"A hybrid method for relation extraction from biomedical literature","year":2005,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Relation (database); Extraction (chemistry); Relationship extraction; Information retrieval; Data mining; Data science; Chromatography; Chemistry","score_opus":0.015439791927503702,"score_gpt":0.3638849635772685,"score_spread":0.3484451716497648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2052547205","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011346042,0.0016871231,0.9645719,0.0004547929,0.00018018665,0.00062646274,0.0063755093,0.0124128265,0.002345117],"genre_scores_gemma":[0.025506318,0.0005284757,0.9584759,0.00013986762,0.00006948623,0.0004154021,0.010764784,0.00035322734,0.0037465177],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99717146,0.00039470496,0.0005587053,0.00063407456,0.0011392969,0.000101839534],"domain_scores_gemma":[0.99482155,0.0029563657,0.000262051,0.0006454765,0.001146822,0.00016776416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022411058,0.0012057069,0.0018681429,0.014901161,0.001349469,0.0034364902,0.0016092216,0.0014332604,0.0072363736],"category_scores_gemma":[0.0061781895,0.0008075724,0.0021748554,0.012070149,0.0004805633,0.0037604533,0.002518992,0.0011225583,0.005430728],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035472997,0.000292403,0.00321801,0.0017215702,0.0005593479,0.00059150567,0.0006235289,0.0017519067,0.037585985,0.0055865436,0.018451007,0.92926335],"study_design_scores_gemma":[0.00070490275,0.0010793145,0.02612705,0.0012408738,0.0031810794,0.013278506,0.002402017,0.3815831,0.13303348,0.060173955,0.37658072,0.00061496283],"about_ca_topic_score_codex":0.002856553,"about_ca_topic_score_gemma":0.0066695977,"teacher_disagreement_score":0.014901161,"about_ca_system_score_codex":0.00048410925,"about_ca_system_score_gemma":0.0022816132,"threshold_uncertainty_score":0.024208069},"labels":[],"label_agreement":null},{"id":"W2053260987","doi":"10.1152/physiolgenomics.00226.2006","title":"Functional nonsynonymous single nucleotide polymorphisms from the TGF-β protein interaction network","year":2006,"lang":"en","type":"article","venue":"Physiological Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lunenfeld-Tanenbaum Research Institute; University of Toronto; Mount Sinai Hospital","funders":"","keywords":"Nonsynonymous substitution; Biology; Single-nucleotide polymorphism; Genetics; Transforming growth factor; Gene; Computational biology; Genotype; Cell biology; Genome","score_opus":0.026452777855213484,"score_gpt":0.22667147722370556,"score_spread":0.20021869936849207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053260987","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9798749,0.0007174166,0.012460883,0.0001626138,0.000014301097,0.000056516754,0.0058448645,0.00019268737,0.0006758634],"genre_scores_gemma":[0.9672084,0.00037152553,0.015579707,0.0000344226,0.000011177214,0.00008716441,0.01647583,0.00002461372,0.00020717752],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995005,0.00013576556,0.00006264957,0.000116190684,0.00014279844,0.000042062216],"domain_scores_gemma":[0.9991062,0.0005644661,0.00017413803,0.000049022285,0.000059727092,0.000046433863],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044956888,0.00031799165,0.0003793661,0.0018117393,0.00042843426,0.00042880524,0.00031399642,0.00030464176,0.0011769022],"category_scores_gemma":[0.002125663,0.00012601555,0.00047615333,0.002021399,0.00023223311,0.00030006058,0.00030542753,0.0002695722,0.00017314183],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031153245,0.0004980269,0.52716357,0.0022565948,0.0018235379,0.008298126,0.0007865458,0.043084055,0.23502833,0.0055891005,0.00495698,0.1673998],"study_design_scores_gemma":[0.00020617839,0.00026690194,0.77612644,0.00010744712,0.0011003901,0.008896789,0.0003987982,0.15996909,0.029909775,0.013065651,0.0098836655,0.00006885455],"about_ca_topic_score_codex":0.002022057,"about_ca_topic_score_gemma":0.0050166184,"teacher_disagreement_score":0.002022057,"about_ca_system_score_codex":0.00038457668,"about_ca_system_score_gemma":0.00054235244,"threshold_uncertainty_score":0.004020512},"labels":[],"label_agreement":null},{"id":"W2054145206","doi":"10.1145/1148170.1148336","title":"Concept-based biomedical text retrieval","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Information retrieval; Set (abstract data type); Text retrieval; Artificial intelligence","score_opus":0.008717700589871995,"score_gpt":0.25303673321000514,"score_spread":0.24431903262013313,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2054145206","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025917942,0.019717632,0.93298465,0.0024878695,0.0009826141,0.0009521375,0.0021651979,0.0038000748,0.010991833],"genre_scores_gemma":[0.1668136,0.007343759,0.81331784,0.0013455477,0.0008487595,0.0005940369,0.004232341,0.00017486043,0.0053292913],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99742496,0.00061356684,0.00022699712,0.00038799888,0.0012480362,0.000098534576],"domain_scores_gemma":[0.9975333,0.0011210236,0.00020211194,0.0002650966,0.00080177165,0.0000766772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002556288,0.00082868774,0.0015827948,0.006175506,0.0008053396,0.0016864993,0.0019797478,0.0013576652,0.0043725315],"category_scores_gemma":[0.006396083,0.00027370392,0.0010637763,0.0053929966,0.00082454376,0.004431577,0.0015930429,0.0010785625,0.00377965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029191596,0.00030789804,0.00088293845,0.0016420062,0.00023968879,0.0004109418,0.00034474707,0.00895518,0.047093585,0.03593507,0.038027115,0.86586887],"study_design_scores_gemma":[0.00046076835,0.00102371,0.0060622655,0.0006145357,0.00055908284,0.0056073503,0.0011042015,0.4409602,0.105899885,0.21671653,0.22055359,0.00043785956],"about_ca_topic_score_codex":0.0022035167,"about_ca_topic_score_gemma":0.0015886803,"teacher_disagreement_score":0.006175506,"about_ca_system_score_codex":0.001038685,"about_ca_system_score_gemma":0.0013172402,"threshold_uncertainty_score":0.014627576},"labels":[],"label_agreement":null},{"id":"W2054323589","doi":"10.3166/ria.25.473-497","title":"évaluer la conformité des prescriptions médicamenteuses aux recommandations de pratique thérapeutique. Utilisation d'un raisonnement ontologique en OWL2","year":2011,"lang":"fr","type":"article","venue":"Revue d intelligence artificielle","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Humanities; Medicine; Philosophy","score_opus":0.07851999518362095,"score_gpt":0.3161203095239581,"score_spread":0.23760031434033713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2054323589","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.83427626,0.002735461,0.1372236,0.0020102814,0.0002986148,0.000744984,0.0089920135,0.005683201,0.008035583],"genre_scores_gemma":[0.8665745,0.0005968948,0.1259788,0.00024215184,0.000043907727,0.00030659125,0.0039366684,0.00019995435,0.0021204466],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99269533,0.0029633944,0.000996926,0.000999986,0.002154646,0.00018982678],"domain_scores_gemma":[0.9570372,0.033970594,0.0029269878,0.0015241799,0.004158335,0.0003826403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0089058895,0.0006432799,0.00066014857,0.0028925969,0.00039296807,0.0021945522,0.0006279073,0.0011097651,0.0020683333],"category_scores_gemma":[0.060731255,0.00028816846,0.00094899326,0.0018268913,0.00035023087,0.0014684944,0.0007151162,0.0009965108,0.0006489884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00326375,0.0013469222,0.21027884,0.0013064597,0.0010198441,0.0006393535,0.0023554575,0.0614522,0.023851642,0.0024143588,0.0073190653,0.6847521],"study_design_scores_gemma":[0.00040099514,0.0015554172,0.25308418,0.00046372478,0.00074131315,0.0015057566,0.0023692923,0.65214574,0.05533736,0.007269628,0.02491271,0.0002139366],"about_ca_topic_score_codex":0.018663296,"about_ca_topic_score_gemma":0.015169055,"teacher_disagreement_score":0.018663296,"about_ca_system_score_codex":0.0011717683,"about_ca_system_score_gemma":0.0021989348,"threshold_uncertainty_score":0.04709941},"labels":[],"label_agreement":null},{"id":"W2055205760","doi":"10.2196/resprot.2315","title":"The SADI Personal Health Lens: A Web Browser-Based System for Identifying Personally Relevant Drug Interactions","year":2013,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Carleton University; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Universidad Politécnica de Madrid; Canarie; Microsoft Research; Heart and Stroke Foundation of Canada","keywords":"World Wide Web; Computer science; Web page; Personally identifiable information; JavaScript; Workflow; Web application; Plug-in; Internet privacy; Database; Computer security","score_opus":0.2335661848216555,"score_gpt":0.5180854587097071,"score_spread":0.28451927388805165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2055205760","genre_codex":"software","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035349905,0.0016059404,0.21648096,0.0013347148,0.00019118954,0.0019681498,0.044032007,0.6690782,0.029958954],"genre_scores_gemma":[0.30985272,0.0022316326,0.5671874,0.0027127427,0.00032046955,0.0018104555,0.06632612,0.012746125,0.03681225],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99935883,0.00012132253,0.00007036854,0.00017377557,0.00023317563,0.00004251813],"domain_scores_gemma":[0.9979862,0.00096761936,0.00020231461,0.00028928483,0.0002694264,0.00028508957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017316616,0.0012316909,0.00074335246,0.0029063744,0.0003767889,0.0017240174,0.0012247622,0.000884944,0.020687368],"category_scores_gemma":[0.003992906,0.0005898752,0.0005619036,0.0009121972,0.00037522623,0.0025810932,0.0023587528,0.00080889364,0.007816925],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0044309283,0.0008470449,0.02624218,0.002316634,0.00042644117,0.0015061419,0.0029453246,0.0021083218,0.043581557,0.009550066,0.38921794,0.51682746],"study_design_scores_gemma":[0.0015464022,0.0010258843,0.037982147,0.00089807453,0.0006023215,0.003705874,0.0012785376,0.16402434,0.063650966,0.023158342,0.7014209,0.0007061782],"about_ca_topic_score_codex":0.0031431033,"about_ca_topic_score_gemma":0.005427152,"teacher_disagreement_score":0.020687368,"about_ca_system_score_codex":0.00074283587,"about_ca_system_score_gemma":0.0014557844,"threshold_uncertainty_score":0.06920618},"labels":[],"label_agreement":null},{"id":"W2056088256","doi":"10.1016/s1063-4584(10)60411-6","title":"384 COMBINING CHONDROCYTE GENE EXPRESSION, LITERATURE MINING AND PATHWAY/NETWORK ANALYSIS TO EXTRACT BIOLOGICAL INSIGHTS FROM SMALL-SCALE MICROARRAY DATA","year":2010,"lang":"en","type":"article","venue":"Osteoarthritis and Cartilage","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Chondrocyte; Microarray analysis techniques; Scale (ratio); Microarray; Computational biology; Data mining; Gene expression; Biology; Gene; Computer science; Genetics; Cartilage; Geography; Cartography; Anatomy","score_opus":0.013610907260875287,"score_gpt":0.2335752732577691,"score_spread":0.2199643659968938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2056088256","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22888523,0.010596023,0.6402869,0.0020746272,0.00022522989,0.0011330321,0.096355595,0.012032722,0.008410614],"genre_scores_gemma":[0.19104235,0.0041921926,0.72450286,0.00028176967,0.00010063454,0.00072472356,0.07648312,0.00033087342,0.0023413948],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9992729,0.00012420978,0.00012980278,0.0001982948,0.00023436408,0.000040445946],"domain_scores_gemma":[0.99857724,0.0007695876,0.00020045433,0.00013891331,0.00023869923,0.000075074524],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001173041,0.00094193994,0.0009924467,0.010153078,0.00066755235,0.0014219169,0.0005256342,0.00039128683,0.0018231597],"category_scores_gemma":[0.0024220874,0.00025801625,0.0015517569,0.007671793,0.0003046159,0.0010201694,0.00075166905,0.00055330945,0.001328889],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006651166,0.00060600776,0.031074755,0.005466325,0.0018284062,0.0016717003,0.00060254487,0.011175539,0.26589963,0.003854537,0.010561039,0.6665944],"study_design_scores_gemma":[0.00031787116,0.0010205605,0.20885462,0.0012117445,0.006097276,0.004341769,0.0021123767,0.29329607,0.23649246,0.07344318,0.17241427,0.0003978441],"about_ca_topic_score_codex":0.0035689673,"about_ca_topic_score_gemma":0.007114136,"teacher_disagreement_score":0.010153078,"about_ca_system_score_codex":0.00057490176,"about_ca_system_score_gemma":0.0017719411,"threshold_uncertainty_score":0.00709641},"labels":[],"label_agreement":null},{"id":"W2057040657","doi":"10.3166/ria.25.445-472","title":"Intelligence artificielle, ontologies et connaissances en médecine Les limites de la mécanisation de la pensée","year":2011,"lang":"fr","type":"article","venue":"Revue d intelligence artificielle","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Humanities; Philosophy","score_opus":0.07135101080285823,"score_gpt":0.3296922694418722,"score_spread":0.25834125863901397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057040657","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060403526,0.102541916,0.33556086,0.31208995,0.0026070396,0.00016978096,0.0007189369,0.0006009046,0.18530706],"genre_scores_gemma":[0.7509669,0.042006742,0.15955488,0.01435959,0.0034021155,0.0003123489,0.00052497897,0.00019725917,0.028675199],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98841417,0.0067059225,0.0006933127,0.001021253,0.0028130547,0.00035228048],"domain_scores_gemma":[0.96395695,0.029681278,0.0012695066,0.0019141276,0.002492166,0.00068598456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016062032,0.000624401,0.0011804062,0.0026964468,0.0018919774,0.014993916,0.0015199374,0.0040026344,0.009013479],"category_scores_gemma":[0.040341783,0.0004011566,0.00078479387,0.0034615956,0.011512853,0.010691689,0.0035272462,0.0045276755,0.0010036075],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008077055,0.000058369133,0.0028449558,0.0007575191,0.00007576664,0.00018056679,0.0024715401,0.008245478,0.0010269902,0.79763436,0.015292237,0.17133145],"study_design_scores_gemma":[0.000025823994,0.000036407215,0.0023811583,0.0006472959,0.000030555817,0.0003443439,0.0016744668,0.015718019,0.00060027326,0.8561522,0.12234208,0.00004740989],"about_ca_topic_score_codex":0.007270221,"about_ca_topic_score_gemma":0.0039536096,"teacher_disagreement_score":0.016062032,"about_ca_system_score_codex":0.0044102115,"about_ca_system_score_gemma":0.0034895423,"threshold_uncertainty_score":0.0849452},"labels":[],"label_agreement":null},{"id":"W2057634750","doi":"10.9708/jksci.2012.17.2.139","title":"Probabilistic filtering for a biological knowledge discovery system with text mining and automatic inference","year":2012,"lang":"en","type":"article","venue":"Journal of the Korea Society of Computer and Information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Computer science; Probabilistic logic; False positive paradox; Data mining; Inference; Event (particle physics); Knowledge extraction; Set (abstract data type); Information extraction; False positives and false negatives; Artificial intelligence; Information retrieval; Machine learning","score_opus":0.016813529819517413,"score_gpt":0.24862352958027867,"score_spread":0.23180999976076125,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057634750","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037756737,0.00011472235,0.9792747,0.0004300679,0.000041348125,0.00018400786,0.0002791072,0.0151432725,0.0007571881],"genre_scores_gemma":[0.05590694,0.00013135576,0.93974274,0.0005020985,0.00012508618,0.00040874004,0.00091077894,0.00036614278,0.0019060073],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960847,0.0006982842,0.0005629507,0.0011819865,0.0012726872,0.00019936753],"domain_scores_gemma":[0.9876816,0.008431395,0.00084629003,0.0011215295,0.0016564269,0.0002627505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00891812,0.0012016366,0.0017661175,0.0048193187,0.0024107127,0.0043371073,0.0032196858,0.0029095514,0.0053097545],"category_scores_gemma":[0.014665108,0.0012887482,0.0039259493,0.0031271025,0.0012867704,0.004950867,0.0012918377,0.0019726136,0.002976419],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014566233,0.00084955414,0.012508538,0.0013933252,0.00085605273,0.0013822081,0.0014211516,0.10479836,0.048905324,0.07034513,0.02324389,0.7328398],"study_design_scores_gemma":[0.000119189135,0.00016769423,0.0018024308,0.00008402495,0.00026762165,0.0004248521,0.00005869457,0.9109401,0.020712813,0.044421908,0.020852633,0.00014805542],"about_ca_topic_score_codex":0.010728202,"about_ca_topic_score_gemma":0.010446333,"teacher_disagreement_score":0.010728202,"about_ca_system_score_codex":0.0024807774,"about_ca_system_score_gemma":0.0035418961,"threshold_uncertainty_score":0.047164083},"labels":[],"label_agreement":null},{"id":"W2057663593","doi":"10.1109/bibm.2013.6732511","title":"Patient information extraction in noisy tele-health texts","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Noise (video); Natural language processing; Spelling; Noise reduction; Speech recognition; Information retrieval; Sentence; Artificial intelligence; Information extraction; Linguistics","score_opus":0.00748155273060243,"score_gpt":0.2616006917107492,"score_spread":0.25411913898014676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2057663593","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13295822,0.0017078107,0.85106534,0.0016260982,0.00016942635,0.000497787,0.005543493,0.004338597,0.002093226],"genre_scores_gemma":[0.29992867,0.0012607567,0.67942494,0.00048087153,0.0003417091,0.0004770557,0.014635788,0.0003990765,0.0030510097],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665797,0.0008804458,0.00049745565,0.0007656014,0.0010375156,0.00016106278],"domain_scores_gemma":[0.9877055,0.008643475,0.0012908146,0.0007861066,0.0014369555,0.00013711936],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024962896,0.0015372048,0.0014036795,0.005668104,0.0008477014,0.0016189859,0.0012246975,0.0017815153,0.0014483209],"category_scores_gemma":[0.014978739,0.000592938,0.00088890066,0.0051458688,0.00065475435,0.003054153,0.0016530228,0.0011499653,0.001678833],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010242466,0.0003985426,0.014146092,0.002042854,0.00019941872,0.0031270885,0.0034886575,0.025052873,0.08164984,0.005648624,0.0110383,0.8521835],"study_design_scores_gemma":[0.00020201162,0.00047152807,0.029770395,0.00037214067,0.00056902616,0.004394197,0.003329386,0.6047775,0.24439947,0.041209448,0.07025871,0.00024608316],"about_ca_topic_score_codex":0.0016197868,"about_ca_topic_score_gemma":0.0022416327,"teacher_disagreement_score":0.005668104,"about_ca_system_score_codex":0.0007974078,"about_ca_system_score_gemma":0.0013867836,"threshold_uncertainty_score":0.013201773},"labels":[],"label_agreement":null},{"id":"W2058038763","doi":"10.1016/j.ymeth.2015.01.014","title":"Text mining of biomedical literature: Doing well, but we could be doing better","year":2015,"lang":"en","type":"editorial","venue":"Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"European Commission","keywords":"Data science; Computational biology; Computer science; Biology","score_opus":0.027114323731289688,"score_gpt":0.3779476610352015,"score_spread":0.3508333373039118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2058038763","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00004982128,0.026601626,0.0012513242,0.14961961,0.8216586,0.00003148662,0.0001594144,0.00012726798,0.0005007185],"genre_scores_gemma":[0.000441749,0.02222692,0.002109911,0.06540619,0.9050697,0.00005761015,0.00017019601,0.00010901429,0.004408677],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9787926,0.0051486595,0.004917915,0.0016841079,0.008953464,0.000503316],"domain_scores_gemma":[0.816958,0.097573414,0.00833752,0.0044242926,0.06393652,0.008770184],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.033437442,0.0026482604,0.0043430193,0.010121656,0.0029221117,0.013043668,0.004422693,0.011896292,0.006072006],"category_scores_gemma":[0.09939679,0.0012970391,0.004132859,0.004742271,0.0060973885,0.013317664,0.0032737306,0.026439834,0.0062354975],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036003807,0.000014597396,0.000047733516,0.000829201,0.000103921084,0.00007776355,0.000038671194,0.000024269935,0.00011518614,0.00074498926,0.97969556,0.018272167],"study_design_scores_gemma":[0.000077731056,0.000028208788,0.0004002933,0.0023169764,0.00021895787,0.00039407887,0.00010688448,0.0002726428,0.00018915605,0.0063776188,0.98955953,0.000057880723],"about_ca_topic_score_codex":0.0025245033,"about_ca_topic_score_gemma":0.007472414,"teacher_disagreement_score":0.96656257,"about_ca_system_score_codex":0.0035710316,"about_ca_system_score_gemma":0.008029632,"threshold_uncertainty_score":0.17683625},"labels":[],"label_agreement":null},{"id":"W2059929112","doi":"10.1007/s11192-008-2139-z","title":"Increasing dominance of English in publications archived by PubMed","year":2009,"lang":"en","type":"article","venue":"Scientometrics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Dominance (genetics); English language; Period (music); Civilization; History; Library science; Linguistics; Computer science; Biology; Art; Philosophy; Archaeology","score_opus":0.01977561581092457,"score_gpt":0.28266987312396274,"score_spread":0.26289425731303817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2059929112","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5235227,0.26825193,0.006579132,0.024636267,0.005057092,0.00021749047,0.100948535,0.0013852878,0.06940151],"genre_scores_gemma":[0.8059564,0.11360193,0.007175595,0.0050023687,0.006931166,0.00023434933,0.046898305,0.0009190815,0.013280928],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9741061,0.005899411,0.006944477,0.003054707,0.0079840915,0.0020112493],"domain_scores_gemma":[0.7292669,0.13963996,0.06066059,0.010451072,0.04528766,0.014693864],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.016784223,0.0011680017,0.002630563,0.11279484,0.0026678094,0.012051634,0.0012253235,0.0011507762,0.0103267],"category_scores_gemma":[0.11534358,0.00046120456,0.001551502,0.14036903,0.0025738275,0.0062984135,0.00727833,0.0011208961,0.00333288],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036842532,0.00027033876,0.38881433,0.038355224,0.004215888,0.0036982144,0.018231587,0.00051131303,0.029511789,0.013539116,0.10090117,0.39826673],"study_design_scores_gemma":[0.0002091203,0.00030627323,0.6054264,0.00711397,0.003433963,0.006967889,0.011485376,0.00048634215,0.0089189485,0.008759787,0.34671518,0.0001767764],"about_ca_topic_score_codex":0.00320317,"about_ca_topic_score_gemma":0.0056414353,"teacher_disagreement_score":0.98321575,"about_ca_system_score_codex":0.001660208,"about_ca_system_score_gemma":0.0057140696,"threshold_uncertainty_score":0.08876449},"labels":[],"label_agreement":null},{"id":"W2061231608","doi":"10.1093/nar/gkq907","title":"The Protein Ontology: a structured representation of protein forms and complexes","year":2010,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":147,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Institute of General Medical Sciences; National Human Genome Research Institute; National Institutes of Health","keywords":"Biology; Ontology; Representation (politics); Computational biology; Open Biomedical Ontologies; Process ontology; Bioinformatics; Computer science; Information retrieval; Suggested Upper Merged Ontology; Semantic Web","score_opus":0.03836785129297262,"score_gpt":0.3623107203154941,"score_spread":0.3239428690225215,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061231608","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008969501,0.0025192196,0.847963,0.0031694323,0.00052727363,0.0006793585,0.09361858,0.016154354,0.026399238],"genre_scores_gemma":[0.06972498,0.005717161,0.77866375,0.0013904995,0.0002724041,0.0010814195,0.12855932,0.0021416252,0.012448754],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99853206,0.00027276235,0.0002556039,0.00031551495,0.0005025049,0.00012150862],"domain_scores_gemma":[0.99843305,0.00048483547,0.00031468648,0.00039633398,0.00024714042,0.00012398283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017257665,0.0012692551,0.0011215893,0.006248281,0.0015941178,0.004277591,0.002289648,0.0016333106,0.007026239],"category_scores_gemma":[0.004202517,0.0007740544,0.00208787,0.00750241,0.0014366036,0.00957847,0.0029643746,0.0023783988,0.0033394245],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002895087,0.00017767915,0.0025839165,0.002101174,0.00023119044,0.00094115426,0.0023848962,0.012419393,0.012105887,0.64102095,0.1619395,0.16380472],"study_design_scores_gemma":[0.000070727525,0.00005565626,0.0014704166,0.00041113194,0.000114126175,0.0009280547,0.00048830087,0.020051785,0.0034084618,0.26399556,0.7089177,0.00008812193],"about_ca_topic_score_codex":0.010495725,"about_ca_topic_score_gemma":0.010973951,"teacher_disagreement_score":0.010495725,"about_ca_system_score_codex":0.0016285084,"about_ca_system_score_gemma":0.0047742706,"threshold_uncertainty_score":0.023505092},"labels":[],"label_agreement":null},{"id":"W2061327015","doi":"10.1089/106652703322756104","title":"Mining the Biomedical Literature in the Genomic Era: An Overview","year":2003,"lang":"en","type":"review","venue":"Journal of Computational Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":264,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Data science; Computer science; Biomedical text mining; sort; Genomics; Genome; Data mining; Text mining; Information retrieval; Biology","score_opus":0.06706577458925418,"score_gpt":0.38961169413093805,"score_spread":0.3225459195416839,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2061327015","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026380452,0.9376332,0.043916147,0.007508941,0.0005460501,0.00010626557,0.00045979396,0.0004016832,0.0067898147],"genre_scores_gemma":[0.0082109645,0.9190524,0.06789776,0.0014365035,0.0007262246,0.000082748666,0.00063938,0.000040910934,0.0019131389],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99865687,0.00037651384,0.00019117889,0.00017609654,0.0005567112,0.000042617598],"domain_scores_gemma":[0.9945779,0.0037484479,0.0003416917,0.00018082824,0.0009720536,0.0001790595],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028242306,0.0005836831,0.0012332721,0.015155518,0.00085394393,0.0032440373,0.0014876598,0.0014163892,0.0023119806],"category_scores_gemma":[0.005903593,0.0004126323,0.00083309814,0.016818259,0.0011709951,0.006688852,0.0010920491,0.0009199887,0.0024074088],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030242316,0.00006137912,0.0010161083,0.007914093,0.00009193821,0.00045879695,0.0003075393,0.0008408768,0.0015208424,0.008288426,0.01664144,0.96282834],"study_design_scores_gemma":[0.000021153977,0.000076894656,0.0041926205,0.008263817,0.00018031607,0.0048652384,0.0009504853,0.0024179416,0.002161783,0.035950508,0.94085735,0.00006192469],"about_ca_topic_score_codex":0.0025795286,"about_ca_topic_score_gemma":0.0043647042,"teacher_disagreement_score":0.015155518,"about_ca_system_score_codex":0.000890231,"about_ca_system_score_gemma":0.002592921,"threshold_uncertainty_score":0.014936149},"labels":[],"label_agreement":null},{"id":"W2062412640","doi":"10.1109/icacci.2013.6637200","title":"Automatic detection of drug interaction mismatches in package inserts","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"The Scarborough Hospital; University of Toronto","funders":"","keywords":"Package insert; Drug; Computer science; Health care; Medicine; Pharmacology; Political science","score_opus":0.00951987461295067,"score_gpt":0.24684962842898683,"score_spread":0.23732975381603616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062412640","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70075357,0.0046936837,0.2301713,0.0017625351,0.00030924132,0.0008840129,0.02775935,0.02666755,0.006998756],"genre_scores_gemma":[0.57662356,0.00092578237,0.3780294,0.0005069239,0.000099094206,0.00025613964,0.04081055,0.0010252143,0.0017232858],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99412054,0.00088150013,0.001146896,0.0012039478,0.0024857703,0.00016140813],"domain_scores_gemma":[0.96883744,0.018077005,0.006417679,0.0021445483,0.004208138,0.0003152136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003123984,0.0010678738,0.0010950441,0.008175817,0.0007364647,0.0016836382,0.0017595658,0.0015717458,0.0016017404],"category_scores_gemma":[0.019921008,0.00063641317,0.0010227462,0.004898948,0.00058164715,0.0027093517,0.0017644899,0.00088637986,0.0007834657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002537106,0.0008872271,0.21147707,0.0053212284,0.00080912246,0.01033011,0.004324344,0.015077054,0.128079,0.012594671,0.038298003,0.57026505],"study_design_scores_gemma":[0.00022163674,0.00069625233,0.14147294,0.0005986247,0.0011627839,0.011036162,0.002681461,0.40181777,0.27121967,0.0124478545,0.15636536,0.00027951456],"about_ca_topic_score_codex":0.0038439715,"about_ca_topic_score_gemma":0.0053919097,"teacher_disagreement_score":0.008175817,"about_ca_system_score_codex":0.0010902398,"about_ca_system_score_gemma":0.0024468165,"threshold_uncertainty_score":0.016521394},"labels":[],"label_agreement":null},{"id":"W2062558911","doi":"10.1186/1756-0500-3-175","title":"OntoFox: web-based support for ontology reuse","year":2010,"lang":"en","type":"article","venue":"BMC Research Notes","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":209,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; BC Cancer Agency","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institute of Allergy and Infectious Diseases; Canadian Institutes of Health Research; Public Health Agency of Canada; National Institutes of Health; Michael Smith Health Research BC; Public Health Agency; University of Michigan","keywords":"Computer science; Ontology; SPARQL; Information retrieval; Open Biomedical Ontologies; RDF; Process ontology; Ontology-based data integration; Interoperability; Suggested Upper Merged Ontology; World Wide Web; Semantic Web","score_opus":0.14242090789211093,"score_gpt":0.4383030431723238,"score_spread":0.2958821352802129,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062558911","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013546433,0.00044862434,0.70752096,0.0009957353,0.00029568377,0.0009122348,0.0057350025,0.24206501,0.028480185],"genre_scores_gemma":[0.12460305,0.0014421586,0.7099069,0.0017690536,0.00026589623,0.0015227821,0.059904926,0.061824247,0.038760994],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99691284,0.0005010315,0.0004442928,0.0004846109,0.0013740828,0.00028316482],"domain_scores_gemma":[0.9923908,0.0024873447,0.0003635122,0.0034226703,0.000806106,0.0005295812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061465134,0.0013837242,0.0008350278,0.002938652,0.001378419,0.004794136,0.0040623555,0.0018148567,0.015355072],"category_scores_gemma":[0.014476969,0.001389258,0.0024007554,0.0021627499,0.0015432871,0.01148101,0.011617215,0.0025711812,0.0073634307],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018321418,0.0007812651,0.0065887854,0.0022653013,0.00043815473,0.004223328,0.0049351486,0.008278245,0.03474569,0.16514917,0.23981641,0.5309463],"study_design_scores_gemma":[0.0003484967,0.00013067915,0.0027823397,0.00050749944,0.00013119537,0.0014519809,0.00061896496,0.049961902,0.021396223,0.10145506,0.8209466,0.00026893063],"about_ca_topic_score_codex":0.0057294616,"about_ca_topic_score_gemma":0.0053103436,"teacher_disagreement_score":0.015355072,"about_ca_system_score_codex":0.0014808183,"about_ca_system_score_gemma":0.0021507926,"threshold_uncertainty_score":0.05136788},"labels":[],"label_agreement":null},{"id":"W2062707524","doi":"10.1371/journal.pone.0115892","title":"Machine Learning for Biomedical Literature Triage","year":2014,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Genome Alberta; Genome Canada","keywords":"Triage; Computer science; Machine learning; Artificial intelligence; Naive Bayes classifier; Support vector machine; Task (project management); Set (abstract data type); Domain (mathematical analysis); Logistic regression; Data mining; Medicine; Engineering; Mathematics","score_opus":0.027492779605644325,"score_gpt":0.25160609640659165,"score_spread":0.22411331680094732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2062707524","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011931536,0.003380655,0.95483303,0.0031066407,0.0003668513,0.0003854751,0.0019497732,0.02130665,0.002739386],"genre_scores_gemma":[0.12565774,0.0011941486,0.86640495,0.0006455537,0.00032788547,0.00043039926,0.0035884038,0.00038065974,0.0013702816],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99403185,0.0021776918,0.0006742,0.0009991054,0.0019355571,0.00018160215],"domain_scores_gemma":[0.9749497,0.015778165,0.0017355137,0.003828391,0.0032266858,0.00048149234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007336524,0.00093265687,0.0013179898,0.005587585,0.0012627695,0.0025810748,0.001967076,0.0016479973,0.004432319],"category_scores_gemma":[0.039921757,0.00044556826,0.0010220634,0.005157851,0.0007746004,0.0035842208,0.001983034,0.0022123018,0.005295652],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002804521,0.00023861778,0.004268201,0.00068051054,0.00013304633,0.00018535623,0.00027313814,0.019122414,0.0047229202,0.0103889685,0.0271538,0.9325526],"study_design_scores_gemma":[0.00012518941,0.00020432296,0.0043241,0.00043005153,0.000118591175,0.00078788423,0.00031229452,0.7737452,0.01602718,0.13947585,0.06432862,0.00012076527],"about_ca_topic_score_codex":0.002394414,"about_ca_topic_score_gemma":0.0025931848,"teacher_disagreement_score":0.007336524,"about_ca_system_score_codex":0.0011284562,"about_ca_system_score_gemma":0.0028294607,"threshold_uncertainty_score":0.038799703},"labels":[],"label_agreement":null},{"id":"W2064003456","doi":"10.2202/1553-3840.1200","title":"Research that Matters: Linking Researchers, Practitioners, Decision-Makers and the Public: Abstracts from the Fifth Annual IN-CAM Research Symposium November 7 to 9, 2008, Toronto, Canada","year":2008,"lang":"en","type":"article","venue":"Journal of Complementary and Integrative Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto; University of Calgary","funders":"Health Canada; Canadian Institutes of Health Research; Sick Kids Foundation; Japan Agency for Medical Research and Development; Canadian Health Services Research Foundation; Advanced Foods and Materials Network","keywords":"Pharmacy; Library science; Medical education; Medicine; Family medicine; Political science","score_opus":0.10561383368798359,"score_gpt":0.39979122386838795,"score_spread":0.29417739018040434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064003456","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025796157,0.13057657,0.0014334486,0.7706748,0.063135184,0.00072853686,0.0032530038,0.00021705986,0.027401775],"genre_scores_gemma":[0.04863955,0.35828432,0.011821677,0.118046,0.09767148,0.0014675854,0.010314054,0.00068399834,0.35307136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99101263,0.0019433814,0.00063755026,0.0006249,0.004288351,0.0014930493],"domain_scores_gemma":[0.9620978,0.0062494385,0.002111347,0.00046586402,0.015793515,0.013281957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027184032,0.0020796817,0.001663237,0.0045440923,0.009612394,0.014895263,0.0016968764,0.005868765,0.045066148],"category_scores_gemma":[0.027853047,0.0011619289,0.0006164091,0.0072760913,0.0040523536,0.005624627,0.005546568,0.0055265226,0.010022854],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027845268,0.0000082262895,0.00014849952,0.0003054379,0.000006991436,0.00005186819,0.001236553,0.000016032855,0.00016277503,0.00028309875,0.98684675,0.010905993],"study_design_scores_gemma":[0.000033742544,0.00001890707,0.004127548,0.0016380905,0.00002905906,0.00007489832,0.007277181,0.000034900546,0.00014973116,0.0011735487,0.98539644,0.000045919587],"about_ca_topic_score_codex":0.2909723,"about_ca_topic_score_gemma":0.72072935,"teacher_disagreement_score":0.2909723,"about_ca_system_score_codex":0.037605103,"about_ca_system_score_gemma":0.08554108,"threshold_uncertainty_score":0.57855725},"labels":[],"label_agreement":null},{"id":"W2064299012","doi":"10.1371/journal.pone.0060954","title":"Approximate Subgraph Matching-Based Literature Mining for Biomedical Events and Relations","year":2013,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Biomedical text mining; Computer science; Subgraph isomorphism problem; Generalizability theory; Matching (statistics); Induced subgraph isomorphism problem; Data mining; Sentence; Relationship extraction; Graph; Information extraction; Natural language processing; Artificial intelligence; Machine learning; Theoretical computer science; Text mining; Mathematics","score_opus":0.022487262789445638,"score_gpt":0.23931997230192176,"score_spread":0.21683270951247613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064299012","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06786042,0.006179886,0.894072,0.0013902682,0.00014014414,0.0009647954,0.01582321,0.008659272,0.0049100406],"genre_scores_gemma":[0.16952182,0.0026597315,0.7973913,0.00021790617,0.00012796318,0.00044284173,0.027988795,0.00023545462,0.0014141916],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975999,0.0005913055,0.00044008417,0.00063506584,0.0006305415,0.0001030701],"domain_scores_gemma":[0.99374986,0.003741361,0.0008642987,0.00075257575,0.00076292607,0.00012900317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019283153,0.0008666777,0.0010894451,0.020276237,0.0010819467,0.0014854118,0.0014757569,0.0010672809,0.002917694],"category_scores_gemma":[0.013978893,0.00044065964,0.0019366781,0.015534706,0.0006201303,0.003382256,0.0018756134,0.0006020287,0.0014041836],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046913323,0.00031004884,0.014775933,0.004826519,0.0008233372,0.0027277966,0.0013916346,0.0330655,0.048397493,0.029983817,0.02177407,0.8414547],"study_design_scores_gemma":[0.00023368928,0.00042304324,0.024421507,0.0007236454,0.0014281564,0.007394872,0.0018013793,0.53994274,0.055382848,0.25784725,0.11021434,0.00018643153],"about_ca_topic_score_codex":0.0035709941,"about_ca_topic_score_gemma":0.007569564,"teacher_disagreement_score":0.020276237,"about_ca_system_score_codex":0.00080710626,"about_ca_system_score_gemma":0.0024700945,"threshold_uncertainty_score":0.010197997},"labels":[],"label_agreement":null},{"id":"W2064957473","doi":"10.1186/gm376","title":"Inferring novel gene-disease associations using Medical Subject Heading Over-representation Profiles","year":2012,"lang":"en","type":"article","venue":"Genome Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Ontario Institute for Cancer Research; Child and Family Research Institute; University of British Columbia","funders":"National Institute of General Medical Sciences; Michael Smith Health Research BC; Canadian Institutes of Health Research; Government of Ontario; Ontario Institute for Cancer Research","keywords":"Disease; Similarity (geometry); MEDLINE; Computer science; Computational biology; Identification (biology); Health informatics; Informatics; Representation (politics); Subject (documents); Gene Annotation; Bioinformatics; Gene; Medicine; Artificial intelligence; Biology; Genetics; Genome; Pathology","score_opus":0.06901671946663247,"score_gpt":0.36249020228543427,"score_spread":0.2934734828188018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2064957473","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.69600874,0.0054849447,0.25914928,0.0015163912,0.00013469029,0.00058403966,0.029483791,0.0026194109,0.005018673],"genre_scores_gemma":[0.8180295,0.0015608673,0.15867625,0.00015997369,0.00018413842,0.0002901004,0.020637095,0.00008019403,0.0003818407],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976914,0.0007777578,0.0004162438,0.0004899032,0.0005268719,0.00009779015],"domain_scores_gemma":[0.97963184,0.014125987,0.0039009894,0.0007979238,0.0011463031,0.00039699383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046345266,0.0007627364,0.00072351005,0.014852237,0.00045163964,0.0022523252,0.00053992844,0.0007554254,0.0019270629],"category_scores_gemma":[0.027721262,0.00015343155,0.0011584442,0.010005688,0.00030809714,0.0016755188,0.001498914,0.0006475062,0.0008567137],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020491167,0.00045508557,0.4631557,0.003948121,0.0014977247,0.001517357,0.0008533115,0.032007568,0.052891713,0.01055671,0.008071228,0.42299634],"study_design_scores_gemma":[0.00035105343,0.0011023938,0.34651807,0.0009776094,0.002173827,0.0041383477,0.0016210331,0.45923123,0.059752777,0.094155304,0.029751701,0.0002266727],"about_ca_topic_score_codex":0.0014988426,"about_ca_topic_score_gemma":0.0017993752,"teacher_disagreement_score":0.014852237,"about_ca_system_score_codex":0.00073521154,"about_ca_system_score_gemma":0.0015410315,"threshold_uncertainty_score":0.024510026},"labels":[],"label_agreement":null},{"id":"W2065107263","doi":"10.1002/pmic.201400320","title":"VennDIS: A JavaFX‐based Venn and Euler diagram software to generate publication quality figures","year":2014,"lang":"en","type":"article","venue":"PROTEOMICS","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":36,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research; University of Toronto; Princess Margaret Cancer Centre; University Health Network","funders":"Canadian Institutes of Health Research","keywords":"Venn diagram; Computer science; Software; Quality (philosophy); Diagram; World Wide Web; Information retrieval; Database; Mathematics; Programming language; Mathematics education; Physics","score_opus":0.021415533077307043,"score_gpt":0.2923740203773332,"score_spread":0.27095848730002614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2065107263","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035925708,0.0009628219,0.6340449,0.0005439464,0.0009185594,0.00095642766,0.0532304,0.29518044,0.010569998],"genre_scores_gemma":[0.014563767,0.0007969768,0.86451024,0.00035024167,0.00020349893,0.0023275851,0.050440475,0.05948798,0.007319293],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966736,0.000977169,0.0005753711,0.00061760633,0.0009972216,0.00015897023],"domain_scores_gemma":[0.98535603,0.009246817,0.0013143636,0.0013328127,0.0022601387,0.00048979354],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010685775,0.0017784651,0.0016571846,0.008480342,0.0010509889,0.0031244163,0.0029916808,0.0012181115,0.08757421],"category_scores_gemma":[0.027050013,0.0018308223,0.0018067154,0.005235029,0.0005171445,0.0035430186,0.0026471727,0.002367376,0.023464533],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012092717,0.00018565964,0.0059318235,0.0068335226,0.0006431127,0.0005886526,0.0012859734,0.008097684,0.017665103,0.028953865,0.5980455,0.3305598],"study_design_scores_gemma":[0.00053624593,0.00018249845,0.0052633584,0.001154677,0.00025489068,0.0010063803,0.00021815018,0.050777115,0.0297495,0.037209556,0.87325555,0.00039201774],"about_ca_topic_score_codex":0.0021726866,"about_ca_topic_score_gemma":0.0029790578,"teacher_disagreement_score":0.08757421,"about_ca_system_score_codex":0.0010181265,"about_ca_system_score_gemma":0.003351865,"threshold_uncertainty_score":0.29296494},"labels":[],"label_agreement":null},{"id":"W2067284187","doi":"10.1186/1472-6947-10-29","title":"Combining classifiers for robust PICO element detection","year":2010,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":140,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Université de Montréal","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Computer science; Data mining; Identification (biology); Task (project management); Health informatics; Information retrieval; Element (criminal law); Process (computing); Population; Artificial intelligence; Machine learning; Pattern recognition (psychology); Medicine","score_opus":0.03433126563978311,"score_gpt":0.3221974987188929,"score_spread":0.2878662330791098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067284187","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12823482,0.0029689088,0.85981596,0.0005941261,0.00033249595,0.00047550036,0.0005913103,0.004153165,0.0028336416],"genre_scores_gemma":[0.5926025,0.0006507878,0.40169126,0.0003188132,0.00048527675,0.00047075227,0.001822685,0.00019135531,0.0017665547],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941943,0.0015436115,0.00074833556,0.001490859,0.0015954963,0.00042736946],"domain_scores_gemma":[0.9844551,0.010019339,0.0009918718,0.0010691314,0.0031818002,0.00028273472],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008270025,0.0019429049,0.0024325172,0.0063249436,0.0009821815,0.0019694825,0.0017295321,0.0025716457,0.0012775041],"category_scores_gemma":[0.020842796,0.00052438665,0.0015149191,0.002571307,0.0006033852,0.0029030987,0.0015265937,0.0016964015,0.0014721209],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007231598,0.00042063097,0.016492762,0.0003506781,0.0004503251,0.00023035434,0.00018530652,0.035419874,0.015417225,0.0012064476,0.0053377897,0.9237655],"study_design_scores_gemma":[0.00006075747,0.0002702183,0.0062195775,0.00008955384,0.00030327818,0.00039944475,0.00013904752,0.9629356,0.019276975,0.0066814395,0.003564333,0.000059768015],"about_ca_topic_score_codex":0.0021233945,"about_ca_topic_score_gemma":0.0014982839,"teacher_disagreement_score":0.99173,"about_ca_system_score_codex":0.00090226176,"about_ca_system_score_gemma":0.0012123504,"threshold_uncertainty_score":0.043736577},"labels":[],"label_agreement":null},{"id":"W206748266","doi":"10.1007/978-3-319-00651-2_7","title":"Automated Phenotype-Genotype Table Understanding","year":2013,"lang":"en","type":"book-chapter","venue":"Studies in computational intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Biomedicine; Table (database); Style (visual arts); Rhetorical question; Computer science; Natural language processing; Division (mathematics); Table of contents; Information retrieval; Artificial intelligence; Linguistics; Literature; Arithmetic; Art; World Wide Web; Data mining; Biology; Mathematics; Bioinformatics; Philosophy","score_opus":0.14103630439130296,"score_gpt":0.3642759690127401,"score_spread":0.22323966462143716,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W206748266","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013775492,0.00063640025,0.96188074,0.0007466288,0.000093514136,0.00013358197,0.006080281,0.010920823,0.005732685],"genre_scores_gemma":[0.1173625,0.0009092828,0.8560595,0.00042629387,0.00007516415,0.00015809,0.017697575,0.0011109563,0.006200747],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994405,0.000119896125,0.000045057324,0.00017315944,0.00018851715,0.00003295735],"domain_scores_gemma":[0.99820316,0.0011428273,0.000091715025,0.00033138748,0.0001953168,0.000035610967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00095612125,0.0010125813,0.0009895765,0.0018065243,0.00045740596,0.0018924266,0.0018506654,0.0008973979,0.010031066],"category_scores_gemma":[0.0035643512,0.00041424623,0.0014464711,0.001739082,0.00044217717,0.002461837,0.0016650828,0.0010556822,0.002634643],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015764653,0.00024749796,0.005666914,0.00055588165,0.00013163497,0.0010441003,0.00028858802,0.046853546,0.020357527,0.049930662,0.0356825,0.8390835],"study_design_scores_gemma":[0.0000636932,0.00006354286,0.0043880777,0.00022426753,0.00015577646,0.0015124321,0.00029279626,0.6109518,0.040168215,0.27402005,0.06809767,0.00006174699],"about_ca_topic_score_codex":0.002419149,"about_ca_topic_score_gemma":0.0040683867,"teacher_disagreement_score":0.010031066,"about_ca_system_score_codex":0.00066673756,"about_ca_system_score_gemma":0.0010728611,"threshold_uncertainty_score":0.033557296},"labels":[],"label_agreement":null},{"id":"W2067819940","doi":"10.1038/npre.2009.3868.1","title":"BioPortal: Ontologies and Integrated Data Resources at the Click of a Mouse","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; University of Victoria; Mayo Clinic","keywords":"Open Biomedical Ontologies; Ontology; Computer science; RDF; Search engine indexing; World Wide Web; Information retrieval; Visualization; Semantic Web; Process ontology; Data mining; Suggested Upper Merged Ontology","score_opus":0.024802137941780853,"score_gpt":0.31359438043742743,"score_spread":0.28879224249564656,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2067819940","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003158429,0.0028837924,0.40226853,0.0073995166,0.0020341326,0.0009196068,0.23114936,0.29860663,0.051579982],"genre_scores_gemma":[0.017955268,0.0035208303,0.28556633,0.004108459,0.0011203927,0.0013655119,0.6074134,0.04850386,0.030445935],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99752325,0.00046517717,0.0003542457,0.00031295692,0.0011250961,0.00021927318],"domain_scores_gemma":[0.99505687,0.0015905518,0.00033392472,0.0015655671,0.0005737044,0.0008794289],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004406296,0.0018877562,0.0016170917,0.006698693,0.0012294662,0.005127019,0.002615613,0.0029784646,0.09100817],"category_scores_gemma":[0.011213844,0.00189327,0.001188412,0.0077320402,0.0010427138,0.009098698,0.010270648,0.0035775765,0.08872893],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002885482,0.00009702927,0.0005107715,0.00083915517,0.00008172515,0.00042451755,0.00021273085,0.00027461007,0.0041444763,0.022058459,0.90900487,0.062063035],"study_design_scores_gemma":[0.00024263847,0.000031165066,0.001451286,0.0005267023,0.000050321734,0.000770738,0.00012114978,0.0020737802,0.004618573,0.06130804,0.9286974,0.000108228756],"about_ca_topic_score_codex":0.0015688617,"about_ca_topic_score_gemma":0.0023857574,"teacher_disagreement_score":0.09100817,"about_ca_system_score_codex":0.0009249317,"about_ca_system_score_gemma":0.0019350465,"threshold_uncertainty_score":0.30445266},"labels":[],"label_agreement":null},{"id":"W2068768700","doi":"10.1089/omi.2006.10.199","title":"Development of FuGO: An Ontology for Functional Genomics Investigations","year":2006,"lang":"en","type":"review","venue":"OMICS A Journal of Integrative Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Terry Fox Research Institute","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institute of General Medical Sciences; Natural Environment Research Council; National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council; National Institutes of Health","keywords":"Ontology; Genomics; Functional genomics; Gene ontology; Data science; Computer science; Computational biology; Knowledge management; Biology; Genome; Genetics; Gene; Epistemology; Philosophy","score_opus":0.09581901281397573,"score_gpt":0.3718821284710488,"score_spread":0.2760631156570731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068768700","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034171317,0.09343299,0.8438473,0.007542367,0.0015340655,0.0012758194,0.0034911789,0.0062738643,0.039185274],"genre_scores_gemma":[0.010489822,0.095308445,0.86628574,0.0029660778,0.00037530437,0.0014169946,0.008739347,0.00068539125,0.013732809],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9982034,0.00040002336,0.00025683883,0.00023827014,0.00079431024,0.00010724958],"domain_scores_gemma":[0.99766564,0.000891653,0.00026013167,0.00029383204,0.0007007136,0.00018803061],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005599425,0.0012551293,0.001914675,0.0071027544,0.0011665028,0.0037589849,0.0033506418,0.0018676025,0.0026034156],"category_scores_gemma":[0.004258253,0.0007486523,0.0017051803,0.0060983393,0.0019703365,0.006735805,0.0023558617,0.004483261,0.0036300372],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010669902,0.0001776367,0.0009038134,0.004360949,0.00015113893,0.00074488885,0.0010840901,0.0026430471,0.012670019,0.19515266,0.04149297,0.7405121],"study_design_scores_gemma":[0.000016325428,0.00003370011,0.00079371216,0.0012041572,0.00005274584,0.0012444943,0.0002155745,0.0025856865,0.00280592,0.03308611,0.9578989,0.00006265157],"about_ca_topic_score_codex":0.0060053947,"about_ca_topic_score_gemma":0.0047794958,"teacher_disagreement_score":0.0071027544,"about_ca_system_score_codex":0.0037667663,"about_ca_system_score_gemma":0.005731222,"threshold_uncertainty_score":0.029612958},"labels":[],"label_agreement":null},{"id":"W2068909183","doi":"10.1002/ase.3","title":"CAVEman: Standardized anatomical context for biomedical data mapping","year":2007,"lang":"en","type":"article","venue":"Anatomical Sciences Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; University of Calgary","funders":"Dalhousie University","keywords":"Computer science; Java; Software; Workstation; Context (archaeology); Visualization; Human–computer interaction; Atlas (anatomy); Data science; Artificial intelligence; Medicine; Anatomy; Programming language; Biology","score_opus":0.04928306862812123,"score_gpt":0.37538678001238823,"score_spread":0.326103711384267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068909183","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017133111,0.00043173763,0.92583823,0.0005426113,0.00025175937,0.0007813655,0.0075985105,0.05610998,0.0067325393],"genre_scores_gemma":[0.019016579,0.0004887452,0.95651406,0.0003133965,0.00006392011,0.0012677171,0.014220417,0.0050441385,0.0030709624],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99576265,0.00095656165,0.00085946947,0.0006570728,0.0015741637,0.0001902193],"domain_scores_gemma":[0.99151355,0.0034561518,0.000503125,0.002611845,0.0014714486,0.0004438712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063751163,0.0012462,0.0010514065,0.0064778365,0.0016001839,0.0057060425,0.0032420747,0.0017438434,0.02386334],"category_scores_gemma":[0.022787383,0.0014688364,0.0018327493,0.0057394234,0.0014180405,0.0056007975,0.0069582774,0.0024166198,0.006651265],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056470366,0.00015036155,0.0025742755,0.0025665583,0.00024262114,0.0011205825,0.0026904535,0.006754156,0.013364415,0.13846393,0.1930554,0.63845253],"study_design_scores_gemma":[0.0001809964,0.00009970598,0.0022955681,0.0006683033,0.00010389102,0.0016484972,0.00044377332,0.030514302,0.014819776,0.06913452,0.8799041,0.0001865505],"about_ca_topic_score_codex":0.0072860415,"about_ca_topic_score_gemma":0.013525072,"teacher_disagreement_score":0.02386334,"about_ca_system_score_codex":0.0014335411,"about_ca_system_score_gemma":0.005600327,"threshold_uncertainty_score":0.079830825},"labels":[],"label_agreement":null},{"id":"W2068951398","doi":"10.1158/1538-7445.am2011-2885","title":"Abstract 2885: The NCI-Nature Pathway Interaction Database: A comprehensive resource for cell signaling information","year":2011,"lang":"en","type":"article","venue":"Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Computer science; World Wide Web; Database; XML; SQL; Metadata; Resource (disambiguation); Information retrieval","score_opus":0.1190194912533565,"score_gpt":0.39290945894353185,"score_spread":0.2738899676901754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2068951398","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00080848427,0.0034989582,0.0057681245,0.0006965723,0.00016049178,0.0002681735,0.96466553,0.010190628,0.013943023],"genre_scores_gemma":[0.0021411714,0.0024502585,0.008381533,0.00027526516,0.00007229177,0.00035720426,0.9819228,0.0011322233,0.0032672717],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983689,0.00024907058,0.00042066418,0.00023595073,0.0005801369,0.00014525071],"domain_scores_gemma":[0.9955237,0.0010775952,0.00047662013,0.000648306,0.0013837914,0.00089008285],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024682425,0.0023345805,0.003266912,0.01164263,0.0010481093,0.004868401,0.0045030457,0.0018155521,0.13836117],"category_scores_gemma":[0.008499859,0.0014092756,0.0010484208,0.017042033,0.00035687792,0.0029845808,0.003222077,0.0020609177,0.103936],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003075508,0.0000632391,0.0010409585,0.003279389,0.00008250136,0.00019547869,0.000078840705,0.00070551713,0.0031183176,0.0029968568,0.9443309,0.043800592],"study_design_scores_gemma":[0.00034049537,0.000058482332,0.0043649334,0.00085025915,0.00013339052,0.0004190202,0.00008900617,0.0016135522,0.0027244892,0.0059245564,0.9833914,0.000090435584],"about_ca_topic_score_codex":0.007550213,"about_ca_topic_score_gemma":0.009087063,"teacher_disagreement_score":0.13836117,"about_ca_system_score_codex":0.001476805,"about_ca_system_score_gemma":0.005923211,"threshold_uncertainty_score":0.46286422},"labels":[],"label_agreement":null},{"id":"W2070748946","doi":"10.3115/1567619.1567638","title":"BioKI:Enzymes","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Set (abstract data type); Information retrieval; Process (computing); World Wide Web; Programming language","score_opus":0.005581675927201015,"score_gpt":0.22819468154465544,"score_spread":0.22261300561745442,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070748946","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030398055,0.0077949744,0.09338462,0.0035024136,0.001048167,0.0009922993,0.37038916,0.46521118,0.054637283],"genre_scores_gemma":[0.01523908,0.011511225,0.16687325,0.0025136487,0.00038405217,0.0017542373,0.7125168,0.045941595,0.043266136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974226,0.00042807017,0.00072418735,0.00045491653,0.0007663312,0.0002038947],"domain_scores_gemma":[0.9935894,0.0017392691,0.0011449832,0.0012519658,0.0014816491,0.0007927875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029107013,0.0036352233,0.0028512892,0.0075439885,0.0012340836,0.0073713395,0.0033784069,0.0021826457,0.066091426],"category_scores_gemma":[0.013018236,0.0018787035,0.0021503945,0.010044082,0.00063844107,0.0071981684,0.004039368,0.0040529943,0.18400359],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011572128,0.00014897308,0.0021015324,0.007934729,0.0002906264,0.00056546985,0.00033990404,0.0006112412,0.010690422,0.0109088635,0.8760342,0.08921675],"study_design_scores_gemma":[0.00022314276,0.000053310127,0.0012997699,0.0004385755,0.000107028005,0.0005147565,0.00007798411,0.000930275,0.00834673,0.005872551,0.9820208,0.00011510145],"about_ca_topic_score_codex":0.0017338993,"about_ca_topic_score_gemma":0.0016589061,"teacher_disagreement_score":0.066091426,"about_ca_system_score_codex":0.0015824473,"about_ca_system_score_gemma":0.004942976,"threshold_uncertainty_score":0.22109783},"labels":[],"label_agreement":null},{"id":"W2070983861","doi":"10.1145/1835449.1835673","title":"A survival modeling approach to biomedical search result diversification using wikipedia","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Novelty; Ranking (information retrieval); Computer science; Diversification (marketing strategy); Relevance (law); Information retrieval; Probabilistic logic; Data science; Artificial intelligence","score_opus":0.07376606669905197,"score_gpt":0.32131710034803806,"score_spread":0.2475510336489861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2070983861","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031112833,0.00070928386,0.9647563,0.0006104224,0.00005137344,0.000064158834,0.0002396239,0.00040650147,0.00204954],"genre_scores_gemma":[0.8128688,0.0011692834,0.17780782,0.0002684154,0.00024513403,0.0003180408,0.000798074,0.00014837507,0.0063760974],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986791,0.00055666897,0.000086563574,0.00023104939,0.00032561345,0.00012102052],"domain_scores_gemma":[0.9929646,0.0051749535,0.00063919637,0.00029878828,0.000742127,0.00018019984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033595467,0.0007541314,0.001008859,0.0035735182,0.0006862172,0.0017253163,0.0016049378,0.0012134104,0.0018295812],"category_scores_gemma":[0.011982994,0.00043326645,0.0015542004,0.0026190656,0.00087349134,0.0028711515,0.0010060924,0.0009944092,0.0005148536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017696896,0.0001340422,0.008811446,0.00021498548,0.0002700076,0.00029975604,0.00046265402,0.7651938,0.0026554395,0.101311065,0.0033369707,0.11713281],"study_design_scores_gemma":[0.00000596152,0.000025856722,0.00048546822,0.000007885411,0.000022272718,0.0000513218,0.000019575073,0.9732065,0.00020033465,0.025373887,0.0005843162,0.00001658181],"about_ca_topic_score_codex":0.009690763,"about_ca_topic_score_gemma":0.008212267,"teacher_disagreement_score":0.009690763,"about_ca_system_score_codex":0.0015858676,"about_ca_system_score_gemma":0.0009138989,"threshold_uncertainty_score":0.019268692},"labels":[],"label_agreement":null},{"id":"W2071404615","doi":"10.1038/npre.2008.1784.2","title":"Suggested actions from the Melbourne HVP Information Seminar","year":2008,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Hôpital Notre-Dame","funders":"","keywords":"Library science; Political science; Computer science","score_opus":0.015536491484254381,"score_gpt":0.28408598256252354,"score_spread":0.26854949107826914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2071404615","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006566147,0.008958057,0.0019888217,0.5361407,0.31182718,0.00064211455,0.0018786833,0.001550455,0.13635738],"genre_scores_gemma":[0.00719562,0.003365579,0.0030729019,0.18995553,0.082656875,0.0006252915,0.0017501756,0.0010402896,0.71033776],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99186265,0.0018649452,0.00052124343,0.0009493703,0.0033181089,0.0014836811],"domain_scores_gemma":[0.9813348,0.0031331377,0.00095499365,0.00068041997,0.0050842073,0.008812488],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.013010452,0.0021631534,0.0012003308,0.0031446067,0.0041011833,0.01126423,0.00415458,0.023947323,0.30391252],"category_scores_gemma":[0.027786102,0.001098382,0.0019293224,0.0016622309,0.0014919286,0.007260557,0.009236197,0.017772056,0.15299821],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022010392,0.000014929153,0.000015352782,0.00004645331,0.0000013633555,0.00008540796,0.000037405178,0.000014910684,0.00008142676,0.0010309096,0.9938235,0.0048261257],"study_design_scores_gemma":[0.000015243522,0.000015338892,0.00017273429,0.00007796264,0.000003168594,0.00005309715,0.00009414876,0.00003510016,0.00006843854,0.0006489305,0.99879944,0.00001637512],"about_ca_topic_score_codex":0.0065094084,"about_ca_topic_score_gemma":0.009542935,"teacher_disagreement_score":0.30391252,"about_ca_system_score_codex":0.0054225978,"about_ca_system_score_gemma":0.008117764,"threshold_uncertainty_score":0.99288434},"labels":[],"label_agreement":null},{"id":"W2072272271","doi":"10.1038/srep01802","title":"InterMOD: integrated data and tools for the unification of model organism research","year":2013,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Human Genome Research Institute; National Institutes of Health; Wellcome Trust","keywords":"Model organism; Organism; DECIPHER; Computational biology; Data science; Consistency (knowledge bases); Function (biology); Unification; Biology; Computer science; Genomics; Genome; Gene; Bioinformatics; Evolutionary biology; Genetics","score_opus":0.1632633676226406,"score_gpt":0.38353508994295593,"score_spread":0.22027172232031533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072272271","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031179944,0.0016399843,0.6938509,0.0020088411,0.0006555162,0.0010542679,0.08823502,0.19972284,0.009714664],"genre_scores_gemma":[0.017713036,0.002016593,0.7100135,0.00093777693,0.00021772101,0.002514283,0.23956162,0.024375867,0.002649561],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9877473,0.0028212767,0.0032336877,0.0019090952,0.0037558826,0.00053273194],"domain_scores_gemma":[0.96517164,0.011943097,0.0030054997,0.0151748555,0.0025486895,0.0021562038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030103924,0.003245007,0.0030411351,0.01608095,0.0021956665,0.010180795,0.008772784,0.0028159919,0.015040029],"category_scores_gemma":[0.04676532,0.0034051973,0.005964695,0.014167231,0.0022383044,0.014543402,0.01823682,0.0064160265,0.009187203],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023274366,0.00094035367,0.010215257,0.0062606037,0.0019335293,0.0021122778,0.0035626679,0.015611144,0.0153760435,0.26316285,0.34362942,0.33486846],"study_design_scores_gemma":[0.0004402357,0.00014637478,0.0040308596,0.0014540801,0.00046788005,0.0007524427,0.00044017844,0.023478668,0.009612656,0.15388452,0.8050104,0.00028165116],"about_ca_topic_score_codex":0.0047269673,"about_ca_topic_score_gemma":0.004523441,"teacher_disagreement_score":0.030103924,"about_ca_system_score_codex":0.0028975345,"about_ca_system_score_gemma":0.0076018646,"threshold_uncertainty_score":0.15920669},"labels":[],"label_agreement":null},{"id":"W2072750870","doi":"10.1038/npre.2010.5443.1","title":"Keynote: A renaissance for the point mutation: from legacy data to semantic web service","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Mutation; Information retrieval; Annotation; Population; World Wide Web; Visualization; Conceptualization; Artificial intelligence; Genetics; Biology","score_opus":0.03738788146197312,"score_gpt":0.33241611869051185,"score_spread":0.2950282372285387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2072750870","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064522233,0.004684828,0.6097877,0.25688592,0.059772015,0.00030899188,0.0025492292,0.0136027755,0.045956288],"genre_scores_gemma":[0.14874554,0.015619364,0.43545952,0.06981224,0.03856368,0.00081656204,0.008737625,0.016407494,0.26583797],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9934255,0.0016998227,0.00043619642,0.0010537588,0.0028274332,0.0005573531],"domain_scores_gemma":[0.98631465,0.004206746,0.00038014408,0.0032229782,0.0038666294,0.002008868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016318506,0.00087066897,0.00088921754,0.0020236117,0.002436211,0.012146026,0.0026459352,0.0051434333,0.02595453],"category_scores_gemma":[0.021058382,0.00074507087,0.001440534,0.002309572,0.004379153,0.02115207,0.00830853,0.011072902,0.015077209],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040061798,0.000090516056,0.00079351873,0.0003993872,0.000053834417,0.0005889327,0.0015515272,0.0015034348,0.0068796924,0.3927106,0.44283375,0.15219416],"study_design_scores_gemma":[0.000028074393,0.000058898586,0.0002762885,0.00023353388,0.000023032047,0.00027839944,0.00049278396,0.0034505862,0.0035395483,0.08600395,0.9055478,0.00006720063],"about_ca_topic_score_codex":0.0036707385,"about_ca_topic_score_gemma":0.0025338226,"teacher_disagreement_score":0.02595453,"about_ca_system_score_codex":0.0028121155,"about_ca_system_score_gemma":0.0032231184,"threshold_uncertainty_score":0.0868265},"labels":[],"label_agreement":null},{"id":"W2073258315","doi":"10.1109/step.2005.15","title":"Interoperability of Data and Knowledge in Distributed Health Care Systems","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Interoperability; XML; Computer science; Clinical decision support system; Health care; Knowledge management; Knowledge extraction; Semantic interoperability; Data science; Decision support system; Data mining; World Wide Web","score_opus":0.03857763169694575,"score_gpt":0.35121325217742866,"score_spread":0.3126356204804829,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073258315","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029909374,0.0015689695,0.94863,0.007823401,0.0001364121,0.00033640748,0.00026717974,0.0011305559,0.010197618],"genre_scores_gemma":[0.5379955,0.001541887,0.45047355,0.0017697422,0.0002829951,0.00062654883,0.0016573124,0.0003282881,0.005324177],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9685888,0.011895555,0.0041899052,0.0041350774,0.00950052,0.001690122],"domain_scores_gemma":[0.9673263,0.015102037,0.0019783678,0.0122091435,0.0025037525,0.00088056765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02994725,0.0006156581,0.0015075346,0.0045137913,0.002588553,0.014581874,0.0049957857,0.0036282006,0.001763119],"category_scores_gemma":[0.039894536,0.0011563088,0.0015570555,0.0058530206,0.0064734,0.019617101,0.014707903,0.0037019341,0.0006858086],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014651354,0.00021836595,0.007032231,0.00051295187,0.00040549977,0.0016824193,0.00552987,0.04374333,0.0029274416,0.75168484,0.0040665967,0.18204999],"study_design_scores_gemma":[0.000093306844,0.000079823585,0.0021654237,0.00041195852,0.00018731345,0.00062095065,0.0018096762,0.08850371,0.0026103263,0.84109753,0.062325962,0.00009400405],"about_ca_topic_score_codex":0.006589697,"about_ca_topic_score_gemma":0.0027384234,"teacher_disagreement_score":0.02994725,"about_ca_system_score_codex":0.0039094426,"about_ca_system_score_gemma":0.0049138316,"threshold_uncertainty_score":0.15837806},"labels":[],"label_agreement":null},{"id":"W2073902761","doi":"10.1016/j.jvir.2008.10.022","title":"The IR Radlex Project: An Interventional Radiology Lexicon—A Collaborative Project of the Radiological Society of North America and the Society of Interventional Radiology","year":2008,"lang":"en","type":"editorial","venue":"Journal of Vascular and Interventional Radiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"","keywords":"Medicine; Terminology; Variety (cybernetics); Lexicon; Radiology; Standardization; Radiology information systems; Medical physics; Computer science; Artificial intelligence; Linguistics","score_opus":0.0163729943652544,"score_gpt":0.3017830878266956,"score_spread":0.2854100934614412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2073902761","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00016134733,0.00806005,0.0024886993,0.08401761,0.9019592,0.0000647574,0.000523231,0.00040137884,0.0023238335],"genre_scores_gemma":[0.0015510647,0.013298881,0.0033854274,0.04681992,0.919054,0.00014706845,0.0009985825,0.00038861617,0.014356464],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9950741,0.001022573,0.00097588147,0.0003455617,0.0024265898,0.0001553306],"domain_scores_gemma":[0.9647343,0.016858028,0.0022911462,0.00067186385,0.013265013,0.0021795558],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009148159,0.0020932695,0.0027573418,0.0060728,0.0014817556,0.006725464,0.0029295832,0.0077221133,0.006315074],"category_scores_gemma":[0.030386863,0.00093689095,0.001727762,0.002508702,0.002183756,0.003982599,0.0016437897,0.015173921,0.0051183454],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028174723,0.000008788863,0.00002297684,0.00021038907,0.000021854616,0.00010071404,0.000015120503,0.000025091183,0.00006861875,0.0002601206,0.99239856,0.0068395934],"study_design_scores_gemma":[0.000073243034,0.000017185313,0.00020573888,0.0004189738,0.000113680486,0.0003372237,0.000040635547,0.00025428893,0.00019463255,0.00093445904,0.99739254,0.000017422928],"about_ca_topic_score_codex":0.0025876593,"about_ca_topic_score_gemma":0.0065315748,"teacher_disagreement_score":0.009148159,"about_ca_system_score_codex":0.0023131103,"about_ca_system_score_gemma":0.0053853057,"threshold_uncertainty_score":0.048380673},"labels":[],"label_agreement":null},{"id":"W2074909362","doi":"10.1186/1471-2105-12-303","title":"Prototype semantic infrastructure for automated small molecule classification and annotation in lipidomics","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada; New Brunswick Innovation Foundation; Canarie","keywords":"Ontology; Computer science; Annotation; Service (business); Information retrieval; Semantic Web; Web service; World Wide Web; Lipidomics; Bioinformatics; Artificial intelligence; Biology","score_opus":0.04138146981705494,"score_gpt":0.2698668574848378,"score_spread":0.22848538766778287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2074909362","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012488677,0.00024505684,0.8577705,0.00078514195,0.00007694563,0.0006887194,0.0034730393,0.11736262,0.007109197],"genre_scores_gemma":[0.1186737,0.0006348386,0.8414647,0.00088442647,0.00007335053,0.0010295813,0.02720618,0.0044147945,0.0056183757],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99811447,0.00025948932,0.00020553231,0.00049904553,0.0007186717,0.00020291144],"domain_scores_gemma":[0.99729615,0.0007683888,0.00028932033,0.0007479028,0.00063398725,0.0002641504],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039990055,0.0011813146,0.0009374723,0.003567716,0.0014989653,0.0034718125,0.0038637305,0.0019301956,0.005435299],"category_scores_gemma":[0.005438249,0.00072933617,0.0016877531,0.002675865,0.0015161219,0.0062379004,0.0047356137,0.0016189744,0.0046771653],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027701175,0.0023865597,0.016133301,0.0027220207,0.0005533328,0.0034937775,0.0025468243,0.037045,0.13747945,0.16056708,0.12886849,0.50543404],"study_design_scores_gemma":[0.000307979,0.00031747573,0.005115752,0.00045874694,0.00025185043,0.0015910557,0.0008651116,0.52727664,0.09275392,0.10558329,0.2652159,0.00026226332],"about_ca_topic_score_codex":0.006989971,"about_ca_topic_score_gemma":0.006957787,"teacher_disagreement_score":0.006989971,"about_ca_system_score_codex":0.0025496387,"about_ca_system_score_gemma":0.0038272187,"threshold_uncertainty_score":0.02114904},"labels":[],"label_agreement":null},{"id":"W2075689200","doi":"10.1051/pmed:2001018","title":"Maîtrise de l’ordinateur et de l’information :une formation intégrée et continueau premier cycle des études médicales","year":2001,"lang":"fr","type":"article","venue":"Pédagogie médicale","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal","funders":"Université de Montréal","keywords":"Humanities; Political science; Art","score_opus":0.027186639059852868,"score_gpt":0.3246918154963747,"score_spread":0.29750517643652186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075689200","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5096797,0.026591003,0.046557218,0.28026447,0.0021334563,0.0052922075,0.0010481025,0.001079618,0.12735426],"genre_scores_gemma":[0.89739805,0.0096042575,0.035080425,0.010947234,0.0006058144,0.0029553557,0.0003094675,0.00015681211,0.042942513],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.95894474,0.02544777,0.0018535181,0.0028986423,0.007899695,0.0029555706],"domain_scores_gemma":[0.87731206,0.054111525,0.016389227,0.008907261,0.019455777,0.023824126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.046994157,0.00055943464,0.0011041805,0.0025231254,0.007432941,0.010977201,0.0030598613,0.00332532,0.012658721],"category_scores_gemma":[0.09779562,0.0011137229,0.0009503029,0.0025239163,0.006954991,0.00891039,0.010537684,0.0045516198,0.0022080957],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085171516,0.0016364038,0.05551945,0.005263611,0.0001874052,0.000475624,0.13384822,0.00060054596,0.002821019,0.040580332,0.03603875,0.7221768],"study_design_scores_gemma":[0.0006630143,0.0031230801,0.26607126,0.009321849,0.00036380164,0.00082347,0.103443645,0.0018728101,0.0063559026,0.033965413,0.573592,0.00040368605],"about_ca_topic_score_codex":0.04537581,"about_ca_topic_score_gemma":0.057696465,"teacher_disagreement_score":0.046994157,"about_ca_system_score_codex":0.025112273,"about_ca_system_score_gemma":0.08371379,"threshold_uncertainty_score":0.24853188},"labels":[],"label_agreement":null},{"id":"W2075694192","doi":"10.1186/1471-2105-13-s1-s2","title":"SPARQL Assist language-neutral query composer","year":2012,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Canarie; Microsoft Research","keywords":"SPARQL; Computer science; Named graph; Semantics (computer science); World Wide Web; Identifier; Information retrieval; Semantic Web; Query language; Context (archaeology); Task (project management); RDF query language; SPARK (programming language); RDF; Web search query; Web query classification; Programming language; Search engine","score_opus":0.022859403368950572,"score_gpt":0.28317180680258563,"score_spread":0.26031240343363504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2075694192","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012118191,0.00020773511,0.67740697,0.000780641,0.0001632945,0.0007055208,0.003930005,0.29044345,0.01424413],"genre_scores_gemma":[0.28100333,0.00058900216,0.6183768,0.0035002453,0.0003633308,0.0012383761,0.022661438,0.04674813,0.025519464],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99496955,0.0011674246,0.0006684251,0.0011629253,0.0017078457,0.0003238031],"domain_scores_gemma":[0.991935,0.0032086428,0.0003642478,0.0024267156,0.0017503063,0.0003151635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00797061,0.0015548436,0.0010362456,0.0010848835,0.00087732525,0.002866116,0.0026150928,0.0012498195,0.025165204],"category_scores_gemma":[0.011081667,0.00079178915,0.0012929554,0.000994721,0.0011785907,0.0041059265,0.0038647566,0.0018578281,0.010411817],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049361815,0.0007104271,0.010558556,0.0027122276,0.00045227964,0.0027818286,0.0046655824,0.012112077,0.10306396,0.10502899,0.34059748,0.41238034],"study_design_scores_gemma":[0.00048426207,0.00030873384,0.0024581112,0.00027232963,0.00019934583,0.002141246,0.00086890836,0.19841826,0.16812646,0.06158101,0.5648383,0.00030309928],"about_ca_topic_score_codex":0.001698491,"about_ca_topic_score_gemma":0.0012987118,"teacher_disagreement_score":0.025165204,"about_ca_system_score_codex":0.0010003686,"about_ca_system_score_gemma":0.0014426317,"threshold_uncertainty_score":0.08418596},"labels":[],"label_agreement":null},{"id":"W2076069439","doi":"10.1186/1472-6947-8-s1-s3","title":"Experiences mapping a legacy interface terminology to SNOMED CT","year":2008,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hewlett-Packard (Canada)","funders":"U.S. National Library of Medicine","keywords":"SNOMED CT; Terminology; Systematized Nomenclature of Medicine; Computer science; Health informatics; Interface (matter); Artificial intelligence; Information retrieval; Medicine; Pathology; Linguistics","score_opus":0.05093479367705277,"score_gpt":0.34140176113099774,"score_spread":0.29046696745394496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2076069439","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7092925,0.001445003,0.24443339,0.008799323,0.0005189937,0.0007836559,0.0007751601,0.002003257,0.03194877],"genre_scores_gemma":[0.7382328,0.0015891197,0.24690695,0.002436305,0.000120849865,0.00035519042,0.0015171339,0.0012183331,0.007623375],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.97075874,0.019830897,0.0020978034,0.0019044422,0.0045568254,0.0008512613],"domain_scores_gemma":[0.9464428,0.032168668,0.0017273239,0.0074293297,0.009938385,0.0022935148],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03742344,0.000962131,0.00051732693,0.0027183017,0.0041645686,0.00490946,0.0035537314,0.0028545612,0.0041724527],"category_scores_gemma":[0.106270894,0.00072409224,0.0010365201,0.0035571319,0.003459573,0.008836783,0.008194104,0.0034317705,0.0012527549],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049509044,0.0020284639,0.03223506,0.002022796,0.00025578192,0.009746339,0.49391368,0.010150435,0.012132476,0.019738324,0.021937145,0.39534444],"study_design_scores_gemma":[0.00027224803,0.002420316,0.022678284,0.0024208897,0.0005216099,0.023059668,0.27957538,0.052330405,0.056824055,0.057550028,0.50164014,0.0007069777],"about_ca_topic_score_codex":0.0075485017,"about_ca_topic_score_gemma":0.008264744,"teacher_disagreement_score":0.03742344,"about_ca_system_score_codex":0.0026732273,"about_ca_system_score_gemma":0.0039151586,"threshold_uncertainty_score":0.19791645},"labels":[],"label_agreement":null},{"id":"W2079137604","doi":"10.1186/1471-2105-11-s6-s14","title":"Discovering gene functional relationships using FAUN (Feature Annotation Using Nonnegative matrix factorization)","year":2010,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; University of British Columbia; Oak Ridge National Laboratory; Eunice Kennedy Shriver National Institute of Child Health and Human Development; University of Memphis","keywords":"Annotation; Non-negative matrix factorization; Computer science; Gene Annotation; Software; Feature (linguistics); Set (abstract data type); Genome; Field (mathematics); Data mining; Information retrieval; Computational biology; Gene; Matrix decomposition; Biology; Artificial intelligence; Genetics; Mathematics","score_opus":0.05097628625394497,"score_gpt":0.2947423372471521,"score_spread":0.24376605099320714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079137604","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016022775,0.00026894495,0.97571564,0.00013907325,0.000030722516,0.00009145985,0.001131991,0.0059174774,0.00068196224],"genre_scores_gemma":[0.064533204,0.00012271313,0.932885,0.00006451633,0.000019963749,0.00015090719,0.0017289119,0.00013522847,0.0003594827],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993144,0.00015826299,0.00006865036,0.00022571711,0.00018363977,0.00004930037],"domain_scores_gemma":[0.99757653,0.0014916065,0.00032727708,0.00019244068,0.00034989865,0.00006235414],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018430094,0.0010103378,0.00070076925,0.0036301163,0.00078188226,0.00089841685,0.0006985852,0.00065373455,0.0031728428],"category_scores_gemma":[0.006362335,0.00031976646,0.001455424,0.0018935936,0.0004784678,0.0013064239,0.0008626156,0.0006479245,0.0008726494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005289541,0.00021486286,0.010701301,0.0007370821,0.00017211975,0.0005471943,0.00068461476,0.041038066,0.052266054,0.010637117,0.0127900755,0.86968243],"study_design_scores_gemma":[0.0000711433,0.00017245578,0.0068348236,0.00009476167,0.00009070395,0.0010511206,0.00016730632,0.9070336,0.022195706,0.045188654,0.016993802,0.000105928935],"about_ca_topic_score_codex":0.0039244615,"about_ca_topic_score_gemma":0.004940661,"teacher_disagreement_score":0.0039244615,"about_ca_system_score_codex":0.0005483718,"about_ca_system_score_gemma":0.00083734526,"threshold_uncertainty_score":0.010614216},"labels":[],"label_agreement":null},{"id":"W2079442270","doi":"10.1002/meet.14504901246","title":"OTO: Ontology term organizer","year":2012,"lang":"en","type":"article","venue":"Proceedings of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"","keywords":"Ontology; Login; Term (time); Computer science; sort; Set (abstract data type); World Wide Web; Information retrieval; Domain (mathematical analysis); Epistemology; Mathematics; Computer security; Philosophy; Physics","score_opus":0.00901649751918087,"score_gpt":0.26520424206366044,"score_spread":0.25618774454447957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079442270","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01536689,0.0006813779,0.4827726,0.0013826112,0.00076316897,0.0028450945,0.27729213,0.20490164,0.013994511],"genre_scores_gemma":[0.053614777,0.0005817026,0.47603732,0.0006454689,0.00025553783,0.0027856298,0.44313222,0.015520642,0.0074266777],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99732924,0.00032165542,0.0006468072,0.00054414134,0.0009001361,0.0002581261],"domain_scores_gemma":[0.99462813,0.0015512091,0.0008458329,0.0013972917,0.0010474428,0.0005300851],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030433277,0.0018331019,0.0011340848,0.008116474,0.0015146631,0.0032404282,0.002071284,0.0010863453,0.016652808],"category_scores_gemma":[0.011010834,0.0010755203,0.0022311956,0.006294938,0.0008217503,0.00546706,0.0036938093,0.002203153,0.009196483],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011121895,0.00034793746,0.00935323,0.0027662986,0.0002712884,0.00097029656,0.0016270002,0.0061060092,0.016354896,0.049354963,0.62716556,0.28457034],"study_design_scores_gemma":[0.00039614033,0.00020429143,0.008879039,0.0007107287,0.00018666073,0.00096594123,0.0009540479,0.053441286,0.023707997,0.06355966,0.84668666,0.00030754204],"about_ca_topic_score_codex":0.008139135,"about_ca_topic_score_gemma":0.00869234,"teacher_disagreement_score":0.016652808,"about_ca_system_score_codex":0.0015996965,"about_ca_system_score_gemma":0.0026059675,"threshold_uncertainty_score":0.055709124},"labels":[],"label_agreement":null},{"id":"W2079585637","doi":"10.1186/1472-6947-10-53","title":"A method for encoding clinical datasets with SNOMED CT","year":2010,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":52,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Grey Nuns Community Hospital; Alberta Health Services; University of Victoria","funders":"Alberta Health Services","keywords":"SNOMED CT; Systematized Nomenclature of Medicine; Terminology; ENCODE; Encoding (memory); Computer science; Set (abstract data type); Vocabulary; Health informatics; Information retrieval; Artificial intelligence; Controlled vocabulary; Data mining; Natural language processing; Medicine; Programming language; Pathology","score_opus":0.054469684966303375,"score_gpt":0.42767014898877614,"score_spread":0.37320046402247276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2079585637","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002062544,0.00022203753,0.98219925,0.0005619305,0.00013470418,0.0004704082,0.0057344353,0.0068391273,0.0017755331],"genre_scores_gemma":[0.009865578,0.0001511328,0.98232657,0.00015337073,0.000027630209,0.00039835644,0.006233007,0.00040324745,0.00044104256],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930401,0.002423902,0.0015891714,0.0008403755,0.0019669684,0.00013947142],"domain_scores_gemma":[0.9824751,0.009885794,0.0011042155,0.0034500402,0.0028458082,0.00023913711],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059995335,0.0014617661,0.00076596683,0.010394963,0.0012500664,0.00408675,0.0022192714,0.0016050785,0.006119889],"category_scores_gemma":[0.03206705,0.000931041,0.0020353603,0.010644659,0.0012501511,0.0044846404,0.0036464632,0.0022737829,0.0029737302],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053871813,0.00017553743,0.0029927774,0.0018310966,0.00029115187,0.0012559716,0.0025790362,0.014605418,0.012896558,0.09737417,0.055974636,0.809485],"study_design_scores_gemma":[0.00031984935,0.0003546039,0.002630284,0.0018845438,0.0003988961,0.0046639172,0.0016841436,0.21184593,0.051991925,0.22656573,0.49716797,0.00049223],"about_ca_topic_score_codex":0.0048387605,"about_ca_topic_score_gemma":0.0069616064,"teacher_disagreement_score":0.010394963,"about_ca_system_score_codex":0.001455045,"about_ca_system_score_gemma":0.004274781,"threshold_uncertainty_score":0.031728983},"labels":[],"label_agreement":null},{"id":"W2080007652","doi":"10.1142/9789812701626_0007","title":"PHENOGO: ASSIGNING PHENOTYPIC CONTEXT TO GENE ONTOLOGY ANNOTATIONS WITH NATURAL LANGUAGE PROCESSING","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":74,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Institutes of Health; Dalhousie University","keywords":"Computer science; Ontology; Context (archaeology); Knowledge base; Biological database; Identifier; Information retrieval; Task (project management); Open Biomedical Ontologies; Artificial intelligence; Natural language processing; Upper ontology; Semantic Web; Bioinformatics; Biology; Suggested Upper Merged Ontology","score_opus":0.00887405322294089,"score_gpt":0.2705938078835546,"score_spread":0.2617197546606137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2080007652","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02426371,0.00061938795,0.7244423,0.001423833,0.00046971106,0.0010154474,0.032742795,0.20643212,0.00859065],"genre_scores_gemma":[0.07047329,0.0005814007,0.87559515,0.000969793,0.00012043439,0.0008113402,0.043249104,0.0052741393,0.0029253755],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991248,0.00015512202,0.0000808639,0.00035064828,0.00024197376,0.000046529414],"domain_scores_gemma":[0.99826956,0.00082319026,0.00025177156,0.0003648197,0.00019852539,0.000092083355],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015381962,0.0014278112,0.00082511286,0.003589432,0.0009061672,0.0017525601,0.0014572956,0.0007357372,0.0063646524],"category_scores_gemma":[0.0047094002,0.00061581953,0.0011121331,0.002152653,0.0008043944,0.0025437362,0.0028943017,0.001130668,0.0021586656],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012216514,0.0004649546,0.01760631,0.0028995397,0.00035979546,0.0024504364,0.001991074,0.013200929,0.05156601,0.024358205,0.17460531,0.70927584],"study_design_scores_gemma":[0.0004998613,0.00034876712,0.024839321,0.00055124296,0.0003736184,0.0015898949,0.0015803826,0.31359068,0.06798692,0.11846578,0.46982336,0.00035025808],"about_ca_topic_score_codex":0.004230553,"about_ca_topic_score_gemma":0.007809735,"teacher_disagreement_score":0.0063646524,"about_ca_system_score_codex":0.00089521555,"about_ca_system_score_gemma":0.0015497177,"threshold_uncertainty_score":0.021291852},"labels":[],"label_agreement":null},{"id":"W2081327572","doi":"10.1145/1882992.1883117","title":"Visual coder","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Terminology; Computer science; Visualization; SNOMED CT; Coding (social sciences); Human–computer interaction; Graphical user interface; Interface (matter); Multidisciplinary approach; Domain (mathematical analysis); Information retrieval; Controlled vocabulary; Data science; Artificial intelligence","score_opus":0.006261683091582631,"score_gpt":0.28595246821112497,"score_spread":0.2796907851195423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081327572","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031610548,0.00032418055,0.50849706,0.0011730718,0.0006279416,0.0014311868,0.04904361,0.3990805,0.03666137],"genre_scores_gemma":[0.06683408,0.0008408899,0.6201343,0.002233371,0.0003609977,0.005024328,0.09307577,0.12887898,0.08261724],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977307,0.00041952013,0.0003070906,0.00048565946,0.0008440423,0.00021309675],"domain_scores_gemma":[0.98970526,0.004067475,0.00040341032,0.0016961838,0.0036468222,0.0004809194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035216543,0.0022914873,0.0012947046,0.0035704868,0.00081694353,0.004630151,0.0033428718,0.0020977585,0.14946367],"category_scores_gemma":[0.019207526,0.0008964215,0.0012090466,0.002082724,0.0009468934,0.0031287426,0.0036441232,0.0023290275,0.05990788],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012797075,0.00014060576,0.0011723787,0.0013656301,0.00005996387,0.0008528533,0.0010179679,0.003016431,0.012331628,0.023628553,0.71085316,0.24428114],"study_design_scores_gemma":[0.00028541987,0.00010615061,0.001466337,0.00053768384,0.000040430834,0.0008365608,0.00035357598,0.03956451,0.026707476,0.02571376,0.9042086,0.0001794649],"about_ca_topic_score_codex":0.0047005047,"about_ca_topic_score_gemma":0.0025007224,"teacher_disagreement_score":0.14946367,"about_ca_system_score_codex":0.0014039583,"about_ca_system_score_gemma":0.0018701536,"threshold_uncertainty_score":0.5000058},"labels":[],"label_agreement":null},{"id":"W2081989102","doi":"10.1145/2508037.2508049","title":"Validation of an ontological medical decision support system for patient treatment using a repository of patient data","year":2013,"lang":"en","type":"article","venue":"ACM Transactions on Intelligent Systems and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Probabilistic logic; Decision support system; Misinformation; Context (archaeology); Artificial intelligence; Data science","score_opus":0.03694217863045635,"score_gpt":0.3027034028506056,"score_spread":0.26576122422014925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2081989102","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21499498,0.00045351742,0.72020775,0.0056195194,0.0005051658,0.0030971372,0.011022061,0.03471539,0.009384553],"genre_scores_gemma":[0.45173115,0.00021938817,0.5319209,0.00080873,0.000046121968,0.0006494264,0.012603,0.00039283425,0.0016285145],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99161,0.0030119028,0.0015343677,0.0012389515,0.0023864568,0.00021839843],"domain_scores_gemma":[0.9716419,0.015371223,0.0011224841,0.0071613253,0.004155769,0.00054727687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018154599,0.0007060795,0.0009636323,0.0029686913,0.0014194443,0.0044640186,0.002124215,0.0015969352,0.0030276794],"category_scores_gemma":[0.05358837,0.000495102,0.0011856441,0.0017720179,0.00084413996,0.0029315993,0.0032849167,0.0014225602,0.0013061689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040602307,0.002602729,0.08716838,0.0018282749,0.0009468057,0.0027750276,0.0043309513,0.11852309,0.022256736,0.046934955,0.029955758,0.6786172],"study_design_scores_gemma":[0.0006258573,0.0004764733,0.014596157,0.0005092646,0.00033139047,0.0007084869,0.0014188392,0.8722434,0.036580473,0.020479,0.05181942,0.00021135027],"about_ca_topic_score_codex":0.011045137,"about_ca_topic_score_gemma":0.010704473,"teacher_disagreement_score":0.018154599,"about_ca_system_score_codex":0.0023643484,"about_ca_system_score_gemma":0.0056782933,"threshold_uncertainty_score":0.09601188},"labels":[],"label_agreement":null},{"id":"W2083515008","doi":"10.7202/002743ar","title":"Some Anatomical and Physiological Aspects of Medical Translation","year":2002,"lang":"en","type":"article","venue":"Meta Journal des traducteurs","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Terminology; Subject (documents); Linguistics; Computer science; Homogeneous; Translation (biology); Field (mathematics); Medical terminology; Scientific terminology; Natural language processing; History; Artificial intelligence; Library science; Biology; Philosophy; Mathematics; Pure mathematics","score_opus":0.0694999105825886,"score_gpt":0.2943331777113231,"score_spread":0.22483326712873453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2083515008","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00791498,0.08034919,0.06735039,0.20786996,0.012306988,0.00024873487,0.0010726722,0.00044942088,0.6224377],"genre_scores_gemma":[0.46562827,0.13835123,0.09306402,0.044800043,0.041081473,0.0011329153,0.001860954,0.000682755,0.21339832],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9950541,0.0026746774,0.0005699786,0.0004225522,0.0009304393,0.00034821234],"domain_scores_gemma":[0.9912504,0.005282245,0.00078477524,0.0008362079,0.00151591,0.00033045138],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052029714,0.00067583995,0.00044636772,0.0052025965,0.0037511059,0.008444861,0.0016399272,0.004009966,0.024814734],"category_scores_gemma":[0.019771062,0.0005034728,0.0007062094,0.006353812,0.018681679,0.009618787,0.002597321,0.003920043,0.0065947813],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000453108,0.000032659776,0.0004137039,0.00048677434,0.000007635907,0.0009949752,0.002522257,0.00019898743,0.0007391983,0.9386192,0.018702753,0.03723653],"study_design_scores_gemma":[0.000013571057,0.000058719936,0.0024184608,0.0005870672,0.000013113865,0.0039120927,0.0024667662,0.00035371137,0.0009190197,0.4988214,0.49040037,0.000035713532],"about_ca_topic_score_codex":0.0018643191,"about_ca_topic_score_gemma":0.0012449578,"teacher_disagreement_score":0.024814734,"about_ca_system_score_codex":0.0024653887,"about_ca_system_score_gemma":0.0025116894,"threshold_uncertainty_score":0.083013535},"labels":[],"label_agreement":null},{"id":"W2084493006","doi":"10.1016/j.ijmedinf.2005.08.008","title":"Amplification of Terminologia anatomica by French language terms using Latin terms matching algorithm: A prototype for other language","year":2005,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Hospitalier Universitaire de Sherbrooke","funders":"Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung","keywords":"Terminology; Computer science; Natural language processing; Linguistics; Information retrieval; Artificial intelligence; Philosophy","score_opus":0.01785284671722043,"score_gpt":0.344074210660259,"score_spread":0.3262213639430386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084493006","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025341293,0.0003637833,0.9473879,0.00052329444,0.00009314979,0.00046901483,0.0019162697,0.018588325,0.0053169085],"genre_scores_gemma":[0.06368615,0.00015412708,0.9267575,0.00017651741,0.000035928955,0.00020629085,0.0040802476,0.0013259528,0.0035772026],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99793893,0.00051262137,0.0002908096,0.0006460106,0.0004368251,0.00017471662],"domain_scores_gemma":[0.996276,0.0011880869,0.00023039825,0.00060183596,0.001557463,0.00014621756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020914834,0.0008067379,0.00095781207,0.0046144845,0.0016153229,0.0027122407,0.001516472,0.0009310321,0.011774214],"category_scores_gemma":[0.0063519035,0.00040291634,0.0022619995,0.0037874542,0.0008227163,0.003544359,0.0021826832,0.0010411801,0.005592592],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067980884,0.00028993975,0.008722223,0.001123511,0.0002203157,0.00084534497,0.003021721,0.0025942274,0.09134483,0.044157367,0.01627522,0.8307255],"study_design_scores_gemma":[0.00034752072,0.00057576824,0.012470835,0.0005674698,0.0010065421,0.0049366686,0.00695312,0.23936363,0.2695945,0.098045714,0.36584926,0.0002889887],"about_ca_topic_score_codex":0.008044485,"about_ca_topic_score_gemma":0.0067514875,"teacher_disagreement_score":0.011774214,"about_ca_system_score_codex":0.0011972968,"about_ca_system_score_gemma":0.0026928692,"threshold_uncertainty_score":0.039388657},"labels":[],"label_agreement":null},{"id":"W2084540039","doi":"10.1109/rose.2013.6698412","title":"Robot ontologies for sensor- and Image-guided surgery","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Robot; Robotics; Artificial intelligence; Computer science; Human–computer interaction; Ontology; Terminology; Service (business); Teleoperation","score_opus":0.03760319281910866,"score_gpt":0.28997756873502273,"score_spread":0.25237437591591405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084540039","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024631289,0.005297426,0.95723236,0.0044024913,0.00058961526,0.00039305698,0.0017320727,0.001477666,0.026412163],"genre_scores_gemma":[0.072027534,0.0074695065,0.9013108,0.0012881254,0.00038494938,0.0007446518,0.0050527155,0.00045217414,0.011269532],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99732566,0.0008521133,0.0005371562,0.00029916668,0.0008401202,0.00014580066],"domain_scores_gemma":[0.997603,0.0010084361,0.00030441434,0.0004980385,0.00044228963,0.00014388241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002836078,0.0007836123,0.00068420917,0.0036569245,0.0019124773,0.004350498,0.0018944714,0.0024943028,0.006674345],"category_scores_gemma":[0.006275788,0.00065682596,0.0018944001,0.004080112,0.0024545174,0.008548036,0.0031202086,0.0023518093,0.0028876637],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030378244,0.0000502553,0.00044162944,0.0008888134,0.000049942504,0.00048117468,0.0009731033,0.0060491925,0.0020109783,0.8470081,0.021452889,0.120563656],"study_design_scores_gemma":[0.000017651408,0.000025030447,0.000815724,0.0008373862,0.00005136669,0.0008887222,0.00057002704,0.02961704,0.0018341525,0.4287849,0.53649896,0.000058967504],"about_ca_topic_score_codex":0.0063530714,"about_ca_topic_score_gemma":0.007680529,"teacher_disagreement_score":0.006674345,"about_ca_system_score_codex":0.002620525,"about_ca_system_score_gemma":0.0040532076,"threshold_uncertainty_score":0.0223279},"labels":[],"label_agreement":null},{"id":"W2084946528","doi":"10.1007/s10791-006-9020-6","title":"Knowledge-based query expansion to support scenario-specific retrieval of medical free text","year":2007,"lang":"en","type":"article","venue":"Information Retrieval","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":79,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; McMaster University","keywords":"Computer science; Query expansion; Unified Medical Language System; Information retrieval; Query language; Testbed; Query optimization; Precision and recall; Exploit; Recall; Data mining; World Wide Web","score_opus":0.020538311628058734,"score_gpt":0.3008048304354492,"score_spread":0.28026651880739045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2084946528","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09532925,0.001011649,0.856819,0.0018722401,0.00019841186,0.0010241999,0.005377531,0.026430693,0.011937067],"genre_scores_gemma":[0.46831605,0.0006388643,0.5130407,0.0006193874,0.00014551415,0.0005836184,0.012729787,0.00063886243,0.00328711],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99880934,0.00043616,0.00015769107,0.00018535949,0.00031923538,0.00009224516],"domain_scores_gemma":[0.9966466,0.002170275,0.00014005954,0.00034771184,0.00058284035,0.00011243212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015976052,0.00064914336,0.00076642825,0.0027620764,0.00048705377,0.0011158787,0.0013263721,0.0011469394,0.00836336],"category_scores_gemma":[0.008567872,0.00031569955,0.00073128066,0.0017691213,0.00032029004,0.0023898704,0.0015273795,0.0006612243,0.00248776],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018154676,0.0009832729,0.006045635,0.0014991778,0.00029644638,0.0028388766,0.001911335,0.06046873,0.08260421,0.02281441,0.09860475,0.72011757],"study_design_scores_gemma":[0.00030968245,0.00021220907,0.0033526153,0.0001591405,0.0002594514,0.0018072606,0.00076803105,0.8807873,0.042344667,0.023469152,0.046420544,0.00010992402],"about_ca_topic_score_codex":0.0031996563,"about_ca_topic_score_gemma":0.004793975,"teacher_disagreement_score":0.00836336,"about_ca_system_score_codex":0.00063796237,"about_ca_system_score_gemma":0.0008615773,"threshold_uncertainty_score":0.027978241},"labels":[],"label_agreement":null},{"id":"W2086185835","doi":"10.1016/s0828-282x(09)70543-4","title":"Acronym Madness – Part 2","year":2009,"lang":"en","type":"article","venue":"Canadian Journal of Cardiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Queen Elizabeth II Health Sciences Centre","funders":"","keywords":"Acronym; Medicine; Meaning (existential); Alternative medicine; Medical education; Family medicine; Linguistics; Pathology","score_opus":0.014346006256433912,"score_gpt":0.24631518589646909,"score_spread":0.23196917964003516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086185835","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008694233,0.018813657,0.11358378,0.020498257,0.2141148,0.0014975528,0.04597846,0.017354745,0.5594645],"genre_scores_gemma":[0.068812795,0.012537424,0.08679852,0.011277732,0.03259173,0.0014428293,0.07513847,0.009827274,0.7015733],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985032,0.0002308019,0.00022886519,0.00029645502,0.0006429601,0.000097697295],"domain_scores_gemma":[0.9946485,0.0010104042,0.00041750327,0.0006048606,0.0028677087,0.00045096633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010613024,0.0012211947,0.0011024709,0.0042369757,0.0011461863,0.0030209322,0.0011918718,0.0010744578,0.21154165],"category_scores_gemma":[0.008237578,0.00035209773,0.0006929804,0.0034230312,0.0009835332,0.0026773189,0.0016708379,0.001543398,0.11058789],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011114384,0.000026533935,0.00027641692,0.00059305923,0.000011651102,0.00015206245,0.00007348955,0.00023411865,0.0031454829,0.011963828,0.8140452,0.1693671],"study_design_scores_gemma":[0.000007947615,0.000021785883,0.0007314341,0.00010074757,0.000005838648,0.00039000387,0.00003084356,0.00027814088,0.000741541,0.0026608487,0.9950199,0.00001092942],"about_ca_topic_score_codex":0.0018780393,"about_ca_topic_score_gemma":0.0017238451,"teacher_disagreement_score":0.21154165,"about_ca_system_score_codex":0.0009289573,"about_ca_system_score_gemma":0.0014754751,"threshold_uncertainty_score":0.70767736},"labels":[],"label_agreement":null},{"id":"W2086676200","doi":"10.1007/s10916-014-0079-0","title":"Developing a Semantic Web Model for Medical Differential Diagnosis Recommendation","year":2014,"lang":"en","type":"article","venue":"Journal of Medical Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":36,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Health informatics; Computer science; Semantic Web; World Wide Web; Differential (mechanical device); Information retrieval; Medicine; Public health; Pathology","score_opus":0.03790750180619841,"score_gpt":0.326552715103138,"score_spread":0.2886452132969396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2086676200","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03172173,0.00061921787,0.95047724,0.0016142437,0.00014577643,0.0004370902,0.0039438372,0.006494567,0.0045463154],"genre_scores_gemma":[0.30272293,0.00076117506,0.6823661,0.00066527264,0.000062943844,0.00040311035,0.008427834,0.00025122825,0.004339436],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991819,0.00015846218,0.0001463378,0.00015063242,0.00030269774,0.000059861635],"domain_scores_gemma":[0.99859804,0.00052975764,0.00007240012,0.0002173953,0.00049511815,0.000087283515],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013621256,0.0004246522,0.0006552728,0.0034799036,0.0008094149,0.0021882951,0.0012017573,0.0013151172,0.002949522],"category_scores_gemma":[0.003797355,0.00040606878,0.0018652537,0.002619741,0.00032187384,0.0036283128,0.0011769213,0.000984526,0.001284997],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074010296,0.0015578449,0.020573694,0.00069072074,0.0008024989,0.0014972992,0.00056206743,0.13830014,0.0131772645,0.12702395,0.036481254,0.6585932],"study_design_scores_gemma":[0.00003759462,0.000037275175,0.0010950022,0.00006495731,0.00012036246,0.00018450712,0.00010969709,0.93818355,0.0037914787,0.043895006,0.012454184,0.00002626661],"about_ca_topic_score_codex":0.0245457,"about_ca_topic_score_gemma":0.03455901,"teacher_disagreement_score":0.0245457,"about_ca_system_score_codex":0.0012700163,"about_ca_system_score_gemma":0.0020133033,"threshold_uncertainty_score":0.048805654},"labels":[],"label_agreement":null},{"id":"W2088628992","doi":"10.1016/j.ymthe.2005.06.263","title":"260. A Database for Managing Neuromuscular Disease Data in the Province of Quebec","year":2005,"lang":"en","type":"article","venue":"Molecular Therapy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Neuromuscular disease; Disease; Database; Biology; Business; Geography; Medicine; Computer science; Internal medicine","score_opus":0.02825970497155946,"score_gpt":0.3035619652217325,"score_spread":0.27530226025017307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2088628992","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03953362,0.0049788603,0.027340483,0.0052571273,0.00036532473,0.0013299631,0.6965347,0.02019504,0.20446487],"genre_scores_gemma":[0.29752564,0.004284853,0.04998374,0.0021432538,0.00018997514,0.0005842329,0.50092167,0.0015531358,0.14281353],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994067,0.00006124721,0.00006946111,0.00010740979,0.00025177412,0.00010341405],"domain_scores_gemma":[0.9960646,0.00032845014,0.00019312913,0.00032240676,0.0025274078,0.00056408934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011728331,0.00041637025,0.00039746804,0.0034661205,0.0021249794,0.0029133637,0.0010640398,0.0005853873,0.04152796],"category_scores_gemma":[0.0050304895,0.00029313768,0.00032015113,0.0058060903,0.00034592513,0.0013347486,0.00072410534,0.00051043305,0.008292449],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045105486,0.00009122392,0.038834017,0.00044954024,0.000092132854,0.0004970984,0.0005627209,0.0022573152,0.0041077123,0.013586552,0.6697435,0.26932713],"study_design_scores_gemma":[0.00015259057,0.00005290314,0.055886846,0.00027698174,0.00007003356,0.0003008398,0.00048153306,0.008154393,0.0019914769,0.001711405,0.93082863,0.00009238789],"about_ca_topic_score_codex":0.9657674,"about_ca_topic_score_gemma":0.96022993,"teacher_disagreement_score":0.04152796,"about_ca_system_score_codex":0.01402322,"about_ca_system_score_gemma":0.029354787,"threshold_uncertainty_score":0.13892484},"labels":[],"label_agreement":null},{"id":"W2089259215","doi":"10.1108/00220411011066763","title":"Classification in a social world: bias and trust","year":2010,"lang":"en","type":"article","venue":"Journal of Documentation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Originality; Computer science; Pluralism (philosophy); Transparency (behavior); Value (mathematics); Foundation (evidence); Trustworthiness; Knowledge management; Epistemology; Sociology; Data science; Political science; Law; Social science; Internet privacy; Computer security","score_opus":0.03093916958424952,"score_gpt":0.3438346007629026,"score_spread":0.3128954311786531,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089259215","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19515249,0.005906758,0.5277025,0.11525331,0.0011066092,0.0005209453,0.00025173122,0.0005438051,0.15356183],"genre_scores_gemma":[0.96534026,0.0008040616,0.030574627,0.0011294815,0.0002545373,0.00017207096,0.000073252435,0.00008067687,0.0015710471],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8003569,0.13309847,0.0145362755,0.008621559,0.04016261,0.0032241894],"domain_scores_gemma":[0.57002217,0.28570443,0.047277134,0.059570923,0.032925796,0.0044995784],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.113621034,0.00055010803,0.0016014667,0.011074531,0.009594497,0.022977564,0.0023855746,0.0034411866,0.003058237],"category_scores_gemma":[0.2594408,0.00074821943,0.00077490124,0.011842737,0.057469875,0.032151103,0.013366665,0.003998736,0.00056654675],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048893853,0.000024816014,0.007975821,0.0002775756,0.00006423428,0.00016523123,0.039707176,0.00086365436,0.00032918167,0.9080439,0.0019436562,0.04055599],"study_design_scores_gemma":[0.000026226462,0.000036838166,0.0022322203,0.0005736373,0.000043506174,0.00022570117,0.015001854,0.0042138128,0.00046482417,0.9409685,0.036152486,0.00006056836],"about_ca_topic_score_codex":0.006178884,"about_ca_topic_score_gemma":0.0032505689,"teacher_disagreement_score":0.9904055,"about_ca_system_score_codex":0.011365732,"about_ca_system_score_gemma":0.010330234,"threshold_uncertainty_score":0.60089266},"labels":[],"label_agreement":null},{"id":"W2089697781","doi":"10.3414/me10-01-0020","title":"Effectiveness of Lexico-syntactic Pattern Matching for Ontology Enrichment with Clinical Documents","year":2010,"lang":"en","type":"article","venue":"Methods of Information in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"National Human Genome Research Institute; National Cancer Institute; University of Pittsburgh","keywords":"Computer science; Matching (statistics); Sentence; Natural language processing; Set (abstract data type); Ontology; Information retrieval; Domain (mathematical analysis); Artificial intelligence; Mathematics; Statistics","score_opus":0.02669870627091075,"score_gpt":0.4329955275027238,"score_spread":0.40629682123181304,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089697781","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7734967,0.0026904962,0.19848275,0.0016667222,0.00014864742,0.0014234331,0.0022202532,0.011087247,0.008783688],"genre_scores_gemma":[0.5494259,0.0007783714,0.44297758,0.00035782193,0.00008429708,0.0003504158,0.0035675594,0.0005175566,0.0019404654],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98587567,0.007370833,0.0016689503,0.0016044575,0.00310874,0.00037126307],"domain_scores_gemma":[0.9174861,0.07028343,0.0032609527,0.0035256802,0.0046887863,0.0007550507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014956412,0.0015725617,0.0010256866,0.0052874447,0.0008548256,0.0017479847,0.0016187446,0.001636137,0.0018520207],"category_scores_gemma":[0.0747392,0.00038081838,0.0011218784,0.003385533,0.0005992274,0.0038252957,0.002026985,0.0007762601,0.0012683572],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0060970956,0.0019747356,0.06441837,0.0030183976,0.0009608365,0.00092405686,0.0016869864,0.0242789,0.05013103,0.001309108,0.0068622157,0.8383384],"study_design_scores_gemma":[0.001277363,0.004950485,0.06712058,0.0006606538,0.0020680814,0.00502964,0.002787455,0.7300282,0.15692222,0.00655422,0.022239415,0.00036174164],"about_ca_topic_score_codex":0.0039001245,"about_ca_topic_score_gemma":0.003938792,"teacher_disagreement_score":0.014956412,"about_ca_system_score_codex":0.0007968153,"about_ca_system_score_gemma":0.0024846995,"threshold_uncertainty_score":0.079098046},"labels":[],"label_agreement":null},{"id":"W2089922777","doi":"10.1109/cjece.2009.5599423","title":"CLASS: a general approach to classifying categorical sequences","year":2009,"lang":"en","type":"article","venue":"Canadian Journal of Electrical and Computer Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Categorical variable; Classifier (UML); Artificial intelligence; Computer science; False positive paradox; Matching (statistics); Class (philosophy); Pattern recognition (psychology); Data mining; Machine learning; Mathematics; Statistics","score_opus":0.010474532208727108,"score_gpt":0.20629269515651957,"score_spread":0.19581816294779247,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2089922777","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007569716,0.00077588914,0.98095393,0.00053946255,0.00024129033,0.0006572347,0.0024957324,0.0037129126,0.0030538747],"genre_scores_gemma":[0.05724614,0.0006754275,0.93252265,0.0003953366,0.00027151583,0.00092144165,0.0044770557,0.00025962782,0.0032307461],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9955449,0.0007492268,0.0006198048,0.0011963393,0.0015962359,0.00029339988],"domain_scores_gemma":[0.9946866,0.0018575437,0.0004912064,0.0012358712,0.0013987286,0.00033001573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034860356,0.001420211,0.0015370643,0.010607176,0.0018318003,0.004165297,0.0039597773,0.0027894687,0.0046838825],"category_scores_gemma":[0.011701499,0.0004497847,0.0023728216,0.008662655,0.0019760225,0.005892895,0.002777683,0.0024474137,0.0026333781],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036156832,0.00034686632,0.013044309,0.0009740128,0.0002394831,0.00025380647,0.0007655131,0.012307928,0.008676919,0.088319086,0.0241848,0.8505257],"study_design_scores_gemma":[0.00010165913,0.0005383941,0.008252207,0.0004763398,0.00020184234,0.0026891788,0.0013790278,0.4179016,0.016331062,0.34630638,0.20557025,0.00025201796],"about_ca_topic_score_codex":0.0042211884,"about_ca_topic_score_gemma":0.003806206,"teacher_disagreement_score":0.010607176,"about_ca_system_score_codex":0.0014675166,"about_ca_system_score_gemma":0.0034982953,"threshold_uncertainty_score":0.018436134},"labels":[],"label_agreement":null},{"id":"W2090242275","doi":"10.1186/1471-2105-9-450","title":"Bluejay 1.0: genome browsing and comparison with rich customization provision and dynamic resource linking","year":2008,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children; University of Calgary","funders":"Genome Canada; Genome Alberta; University of Calgary","keywords":"Personalization; Resource (disambiguation); Computer science; World Wide Web; Genome; DNA microarray; Computational biology; Biology; Data science; Genetics; Gene","score_opus":0.015214036059689802,"score_gpt":0.24463081578491097,"score_spread":0.22941677972522118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2090242275","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30663013,0.001344306,0.092535295,0.0010812739,0.00050044246,0.000839961,0.20153952,0.35975984,0.035769243],"genre_scores_gemma":[0.25733298,0.0005949535,0.2687895,0.00063607394,0.00006183349,0.0007489352,0.41336143,0.0502222,0.008252103],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99892056,0.000153776,0.0000665933,0.00031835973,0.00036323507,0.00017753892],"domain_scores_gemma":[0.99857974,0.00042169017,0.00008239338,0.0003901773,0.00034070824,0.00018530311],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018298642,0.0010644053,0.0005903055,0.0026290282,0.0010581793,0.0018449366,0.0021700684,0.00090608513,0.013708968],"category_scores_gemma":[0.003027521,0.0006893086,0.0010816591,0.0027312448,0.00043283342,0.0020311312,0.0022034706,0.0012576621,0.006035421],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0074799997,0.0018123443,0.034583807,0.0027291423,0.00073068513,0.0017138455,0.0033694208,0.018136956,0.11936879,0.007981828,0.5822759,0.21981724],"study_design_scores_gemma":[0.0016474961,0.0009576742,0.100705266,0.00047956032,0.0003136718,0.0023529257,0.0018789982,0.1353995,0.113390446,0.0076605934,0.63443726,0.0007766621],"about_ca_topic_score_codex":0.02404372,"about_ca_topic_score_gemma":0.021892633,"teacher_disagreement_score":0.02404372,"about_ca_system_score_codex":0.00074475957,"about_ca_system_score_gemma":0.0013224857,"threshold_uncertainty_score":0.047807515},"labels":[],"label_agreement":null},{"id":"W2091810468","doi":"10.1186/1471-2105-9-s3-s1","title":"The Second International Symposium on Languages in Biology and Medicine (LBM) 2007","year":2008,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computational biology; Biology; Data science; Computer science","score_opus":0.01740688202851943,"score_gpt":0.30610580208425237,"score_spread":0.28869892005573294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2091810468","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0108692655,0.19621114,0.12829727,0.17165798,0.28933758,0.0011485154,0.00461393,0.0062799873,0.19158436],"genre_scores_gemma":[0.033844534,0.078254096,0.08283763,0.019477114,0.05606778,0.0012626888,0.008036711,0.0051677227,0.71505165],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9911209,0.0035050274,0.0007173768,0.0012092039,0.0025095234,0.0009379881],"domain_scores_gemma":[0.987795,0.0031029787,0.0005428674,0.0010826743,0.0044209706,0.003055546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013124551,0.0020329268,0.0023808696,0.004958771,0.0020698553,0.012725313,0.0029364675,0.004229848,0.094103895],"category_scores_gemma":[0.01733918,0.0007503247,0.0023809727,0.0022615886,0.0025128268,0.00951257,0.0075120064,0.006457448,0.047818456],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021582769,0.00007452427,0.0002511944,0.0006086495,0.000040323786,0.00014550111,0.0007335385,0.00023552924,0.0021049627,0.020028817,0.77180505,0.20375597],"study_design_scores_gemma":[0.000010896858,0.000046971978,0.00032681192,0.00042511596,0.000013878932,0.00023865375,0.0003073269,0.00037447596,0.00076579204,0.00564968,0.9918133,0.000027107628],"about_ca_topic_score_codex":0.0022875671,"about_ca_topic_score_gemma":0.00355269,"teacher_disagreement_score":0.094103895,"about_ca_system_score_codex":0.005488884,"about_ca_system_score_gemma":0.008080996,"threshold_uncertainty_score":0.3148089},"labels":[],"label_agreement":null},{"id":"W2092375840","doi":"10.12688/f1000research.6140.1","title":"Procedure and datasets to compute links between genes and phenotypes defined by MeSH keywords","year":2015,"lang":"en","type":"preprint","venue":"F1000Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Ottawa Hospital Research Institute; University of Ottawa","keywords":"Computer science; Phenotype; Set (abstract data type); Inference; Ontology; Data mining; Association rule learning; Prioritization; Fuzzy logic; Gene; Computational biology; Biology; Artificial intelligence; Genetics; Programming language","score_opus":0.05195907390469875,"score_gpt":0.35726280425795537,"score_spread":0.30530373035325664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2092375840","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009940541,0.00018148308,0.1459364,0.0014034878,0.0002444833,0.0025633264,0.71392006,0.11501434,0.010795914],"genre_scores_gemma":[0.021746105,0.00014710052,0.27121574,0.00045459036,0.000059045688,0.005457257,0.6912958,0.0048251245,0.004799273],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99832684,0.0002420144,0.0003057908,0.00046404658,0.00053920614,0.00012214991],"domain_scores_gemma":[0.9955433,0.0017259851,0.00023436679,0.0012433649,0.0010647621,0.00018819602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022136709,0.0013239711,0.00071632967,0.0036320311,0.0010058695,0.0019705344,0.0021514774,0.0013868879,0.05082515],"category_scores_gemma":[0.01534452,0.00072764256,0.0014310852,0.0035589882,0.00046780362,0.0015304987,0.0024820112,0.0015240343,0.029031038],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00093509624,0.000275228,0.010183127,0.0014773483,0.00019620788,0.00044322308,0.0002794392,0.009129428,0.009280875,0.017352646,0.8071378,0.14330961],"study_design_scores_gemma":[0.0011213089,0.00020308996,0.01824732,0.00032981075,0.00009413986,0.00075484137,0.00053634593,0.08618016,0.03929056,0.058256935,0.7948213,0.0001640207],"about_ca_topic_score_codex":0.004346854,"about_ca_topic_score_gemma":0.0049071116,"teacher_disagreement_score":0.05082515,"about_ca_system_score_codex":0.0014107507,"about_ca_system_score_gemma":0.003180498,"threshold_uncertainty_score":0.17002708},"labels":[],"label_agreement":null},{"id":"W2095754765","doi":"10.1186/2041-1480-2-s2-s1","title":"The Translational Medicine Ontology and Knowledge Base: driving personalized medicine by bridging the gap between bench and bedside","year":2011,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":91,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"U.S. National Library of Medicine; Engineering and Physical Sciences Research Council; National Institutes of Health","keywords":"SPARQL; Bridging (networking); Computer science; Translational medicine; Knowledge base; Ontology; Precision medicine; Personalized medicine; Translational research; Semantic Web; Data science; Translational science; World Wide Web; RDF; Bioinformatics; Medicine","score_opus":0.040145185584231156,"score_gpt":0.30022983770969225,"score_spread":0.2600846521254611,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095754765","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01010891,0.0019299734,0.92751676,0.0335836,0.00059444027,0.00057571236,0.0033773081,0.004057236,0.018256119],"genre_scores_gemma":[0.05710378,0.0024134086,0.92604965,0.0038238945,0.00026588005,0.00033348036,0.0077193533,0.00041793755,0.0018726268],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943193,0.0022399882,0.00090492575,0.0008392881,0.0014306206,0.0002659901],"domain_scores_gemma":[0.9803577,0.0089421775,0.001476061,0.00436787,0.0033527748,0.0015033983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017792927,0.0006891454,0.0008732421,0.005616419,0.0019962813,0.0063985623,0.0033267585,0.0023043808,0.0036376955],"category_scores_gemma":[0.024533682,0.0006260951,0.0019578217,0.0052688085,0.0036243415,0.011965489,0.006438752,0.0036011296,0.0017750232],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023645452,0.00052709976,0.009035249,0.0023711084,0.0003477639,0.0010763427,0.004070573,0.015646087,0.010842634,0.4244694,0.07834941,0.45302784],"study_design_scores_gemma":[0.00012689897,0.00011355392,0.0026194178,0.002188805,0.00032606194,0.0009473993,0.001656751,0.058931798,0.007961648,0.46724603,0.45773038,0.00015132416],"about_ca_topic_score_codex":0.009028561,"about_ca_topic_score_gemma":0.011142274,"teacher_disagreement_score":0.017792927,"about_ca_system_score_codex":0.0041280575,"about_ca_system_score_gemma":0.016098447,"threshold_uncertainty_score":0.094099104},"labels":[],"label_agreement":null},{"id":"W2095796045","doi":"10.3109/01942638.2013.840463","title":"Are You Knowledgeable About Knowledge Translation?","year":2013,"lang":"en","type":"editorial","venue":"Physical & Occupational Therapy In Pediatrics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; McMaster University","funders":"","keywords":"Knowledge translation; Psychology; Translation (biology); Medical education; Medicine; Physical medicine and rehabilitation; Knowledge management; Computer science","score_opus":0.052683175145104726,"score_gpt":0.3618550675120782,"score_spread":0.30917189236697346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2095796045","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000443181,0.014795134,0.00087148347,0.22405542,0.75839853,0.000023877905,0.00014679926,0.00009895898,0.0015654787],"genre_scores_gemma":[0.0009861728,0.019430835,0.0011353805,0.09834835,0.8705341,0.000071454815,0.00013903067,0.00011071995,0.009243943],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9882212,0.0036903184,0.0021305447,0.00090762775,0.0046067345,0.00044351423],"domain_scores_gemma":[0.87491435,0.088295795,0.005708006,0.002691662,0.024231944,0.004158207],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022380337,0.0019201812,0.0031015244,0.0060381563,0.0031391883,0.011535047,0.003000395,0.019235753,0.01133823],"category_scores_gemma":[0.1098708,0.0012547008,0.0024593743,0.002983592,0.005759474,0.0095644025,0.003074709,0.024802202,0.007440491],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022749518,0.000008273266,0.000021372602,0.00043022117,0.00003711475,0.000100388796,0.000048717284,0.00002051478,0.00003454072,0.0008729898,0.9868795,0.0115236575],"study_design_scores_gemma":[0.000054395354,0.00001357923,0.00016960099,0.0016667772,0.00013723041,0.00023434752,0.00012259257,0.00017647573,0.00012235034,0.005321734,0.9919504,0.000030519022],"about_ca_topic_score_codex":0.0035967128,"about_ca_topic_score_gemma":0.007698572,"teacher_disagreement_score":0.022380337,"about_ca_system_score_codex":0.003512588,"about_ca_system_score_gemma":0.00686928,"threshold_uncertainty_score":0.11835992},"labels":[],"label_agreement":null},{"id":"W2097866676","doi":"10.1016/j.ijmedinf.2004.06.005","title":"Modelling a decision-support system for oncology using rule-based and case-based reasoning methodologies","year":2004,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Guideline; Computer science; Decision support system; Knowledge base; Case-based reasoning; Multidisciplinary approach; Expert system; Medicine; Data mining; Medical physics; Artificial intelligence","score_opus":0.0893818919662332,"score_gpt":0.40512833034550316,"score_spread":0.31574643837926997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2097866676","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049130533,0.0003992817,0.93221414,0.0025374084,0.0001331784,0.00072437304,0.0013374588,0.001802433,0.011721238],"genre_scores_gemma":[0.45107272,0.0005830921,0.54148376,0.00030453014,0.000060111368,0.00045007278,0.0015547076,0.00011531475,0.004375636],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977016,0.00070641504,0.0004056498,0.000305724,0.00065707334,0.00022352321],"domain_scores_gemma":[0.9960295,0.0027592068,0.0003544093,0.00018208203,0.0005082218,0.00016658762],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031430838,0.0010860187,0.0008284912,0.002244766,0.0010871015,0.005801102,0.002949654,0.00268375,0.0060373237],"category_scores_gemma":[0.007820372,0.0007932701,0.0019698855,0.0014108011,0.0011233723,0.0038001956,0.0013970488,0.0011533996,0.0011658203],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004010279,0.0003596252,0.005926311,0.00059712335,0.00027298453,0.005350115,0.0017353004,0.7832687,0.0053214915,0.11934276,0.004404138,0.07302048],"study_design_scores_gemma":[0.00009795982,0.000058597376,0.0004315645,0.00012790255,0.00016149256,0.00053718494,0.00020206776,0.934865,0.0028035447,0.05005366,0.01061189,0.00004918411],"about_ca_topic_score_codex":0.014369478,"about_ca_topic_score_gemma":0.011036176,"teacher_disagreement_score":0.014369478,"about_ca_system_score_codex":0.0019900308,"about_ca_system_score_gemma":0.0025627608,"threshold_uncertainty_score":0.028571725},"labels":[],"label_agreement":null},{"id":"W2099686928","doi":"10.1109/icde.2008.4497616","title":"DescribeX: Interacting with AxPRE Summaries","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Document Structure Description; XML Schema Editor; XML; XML Schema (W3C); Efficient XML Interchange; Schema (genetic algorithms); XML validation; Streaming XML; XML framework; Information retrieval; XML database; Programming language; XML Signature; World Wide Web","score_opus":0.022634840120681003,"score_gpt":0.23492495983224795,"score_spread":0.21229011971156694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2099686928","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011498141,0.0003237944,0.7367111,0.00095307437,0.00012881924,0.00048842886,0.008434125,0.22936703,0.012095461],"genre_scores_gemma":[0.1374563,0.00083954865,0.7401623,0.000997273,0.00015741921,0.0015371314,0.026933424,0.05182248,0.04009414],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99770963,0.00075560407,0.00021905286,0.00038843308,0.0007819003,0.00014530617],"domain_scores_gemma":[0.99232835,0.004744449,0.0003613461,0.0015362848,0.0006955606,0.00033389917],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050536096,0.001212056,0.00084215304,0.002017717,0.00092145463,0.0041099587,0.0020870673,0.0014550546,0.040922236],"category_scores_gemma":[0.017541043,0.001207287,0.0012449708,0.0015792436,0.00081512315,0.008007352,0.0051116794,0.0013566589,0.009065025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038352164,0.00027966485,0.007602068,0.0023248233,0.00022003944,0.0015838625,0.0099843815,0.011627834,0.026020925,0.15317251,0.28136432,0.5019843],"study_design_scores_gemma":[0.0003333094,0.0002760639,0.0017692586,0.0003355144,0.000094485586,0.0009187385,0.0016074317,0.093768805,0.036386542,0.060753584,0.8035524,0.00020381404],"about_ca_topic_score_codex":0.0022115957,"about_ca_topic_score_gemma":0.0025209289,"teacher_disagreement_score":0.040922236,"about_ca_system_score_codex":0.001086142,"about_ca_system_score_gemma":0.0009075505,"threshold_uncertainty_score":0.13689846},"labels":[],"label_agreement":null},{"id":"W2100848433","doi":"10.1002/meet.1450400107","title":"Haystacks and hypotheses","year":2003,"lang":"en","type":"article","venue":"Proceedings of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Haystack; Process (computing); Subject (documents); Computer science; Information retrieval; Term (time); Data science; Scientific literature; Artificial intelligence; Library science; Biology","score_opus":0.010118662018153007,"score_gpt":0.25382810880094414,"score_spread":0.24370944678279113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2100848433","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21136415,0.016254963,0.36178756,0.14353313,0.004017535,0.00093130284,0.0017845249,0.00091441796,0.2594124],"genre_scores_gemma":[0.90941787,0.0038526708,0.06897882,0.00677764,0.0011590918,0.00052535295,0.0010109567,0.00008781538,0.008189699],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98545825,0.007731394,0.00075389876,0.0018713133,0.0036397625,0.00054549286],"domain_scores_gemma":[0.9094602,0.072004125,0.00619526,0.0066450834,0.0042797835,0.0014155351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021973766,0.0010212662,0.0011723428,0.008882204,0.0039034493,0.010924574,0.0027097538,0.003484992,0.028437335],"category_scores_gemma":[0.08361683,0.0007102502,0.001454015,0.0036131435,0.018131757,0.015033263,0.0051821778,0.003473048,0.0018753553],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012806397,0.00006546059,0.004547883,0.00039945368,0.0001294446,0.00074957917,0.004987565,0.0015591672,0.00022739975,0.93849224,0.005672184,0.043041535],"study_design_scores_gemma":[0.000033784818,0.000033138378,0.0010120132,0.00031066933,0.00003243432,0.00029964725,0.0027012783,0.0036487347,0.00018839311,0.9756889,0.016028022,0.000023025868],"about_ca_topic_score_codex":0.0009997735,"about_ca_topic_score_gemma":0.000786897,"teacher_disagreement_score":0.028437335,"about_ca_system_score_codex":0.0028592611,"about_ca_system_score_gemma":0.0021947853,"threshold_uncertainty_score":0.116209805},"labels":[],"label_agreement":null},{"id":"W2102788504","doi":"10.1186/1755-8794-7-s1-s12","title":"Automatic detection and resolution of measurement-unit conflicts in aggregated data","year":2014,"lang":"en","type":"article","venue":"BMC Medical Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital; Health Canada","funders":"Canadian Institutes of Health Research","keywords":"Harmonization; Computer science; Data mining; Unit (ring theory); Units of measurement; Data integration; External Data Representation; Information retrieval; Data science; Artificial intelligence; Mathematics","score_opus":0.07650153793103692,"score_gpt":0.29136898294109764,"score_spread":0.21486744501006072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2102788504","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07329163,0.0006755396,0.9128715,0.0013657062,0.000117137446,0.00048440034,0.0024830264,0.0075305714,0.0011805346],"genre_scores_gemma":[0.25156575,0.00020962226,0.7421405,0.00029983566,0.000061996965,0.00023744164,0.0048379325,0.00044572903,0.00020127976],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97024757,0.0096460935,0.0058726934,0.0054564243,0.007974114,0.00080315484],"domain_scores_gemma":[0.89961225,0.052635122,0.018701104,0.018401835,0.009518757,0.001130987],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030994557,0.0010318874,0.001786595,0.010019935,0.0017944571,0.0056691426,0.0032220017,0.0021199055,0.0007807345],"category_scores_gemma":[0.099109836,0.0007628963,0.002480328,0.010092959,0.0017002361,0.0066930475,0.007129495,0.0026504062,0.0003773139],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001163353,0.00066999707,0.16653223,0.0025035017,0.0014122524,0.002206854,0.00823434,0.044148568,0.035337996,0.056862388,0.017876072,0.66305244],"study_design_scores_gemma":[0.00013348232,0.000184983,0.03741076,0.0006051718,0.00072380784,0.0013338085,0.0033688745,0.67714626,0.07315296,0.1740535,0.031648465,0.00023788014],"about_ca_topic_score_codex":0.003484381,"about_ca_topic_score_gemma":0.0040730108,"teacher_disagreement_score":0.030994557,"about_ca_system_score_codex":0.0019827546,"about_ca_system_score_gemma":0.004416361,"threshold_uncertainty_score":0.16391683},"labels":[],"label_agreement":null},{"id":"W2103153231","doi":"10.1504/ijdmb.2012.048172","title":"Re-ranking with context for high-performance biomedical information retrieval","year":2012,"lang":"en","type":"article","venue":"International Journal of Data Mining and Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"U.S. National Library of Medicine","keywords":"Ranking (information retrieval); Context (archaeology); Information retrieval; Computer science; Data science; Biology; Paleontology","score_opus":0.033992309906928496,"score_gpt":0.30036973770142444,"score_spread":0.26637742779449597,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103153231","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09663832,0.013174506,0.87408113,0.0009583925,0.0005673135,0.0009272879,0.0011069031,0.008806895,0.003739275],"genre_scores_gemma":[0.44412878,0.0019306628,0.54774076,0.0003596257,0.0006506378,0.00035852508,0.0023155569,0.00030801483,0.00220736],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935336,0.0028403695,0.0004550014,0.0007245449,0.0021573193,0.0002890764],"domain_scores_gemma":[0.98792064,0.005457501,0.00079251267,0.0020770428,0.00345053,0.0003017814],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039019126,0.001654339,0.0024693639,0.007882127,0.0013572276,0.0023830032,0.0019381053,0.0012554259,0.0024838191],"category_scores_gemma":[0.019519996,0.00059657067,0.0011014474,0.004752973,0.0005498593,0.0030594144,0.0016489963,0.001448547,0.0022589422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008856748,0.00062733045,0.007416855,0.0009610887,0.00039233288,0.00061840913,0.0004072895,0.04025604,0.048155814,0.0058800937,0.0147647085,0.8796344],"study_design_scores_gemma":[0.00034121805,0.0016436676,0.014180141,0.0002160488,0.0006794527,0.0022817985,0.00059568795,0.8318621,0.07852088,0.040850688,0.028350107,0.00047823577],"about_ca_topic_score_codex":0.0035952863,"about_ca_topic_score_gemma":0.009043274,"teacher_disagreement_score":0.007882127,"about_ca_system_score_codex":0.0008095561,"about_ca_system_score_gemma":0.0016057808,"threshold_uncertainty_score":0.020635545},"labels":[],"label_agreement":null},{"id":"W2103365019","doi":"10.1093/bioinformatics/btq271","title":"SLIMS—a user-friendly sample operations and inventory management system for genotyping labs","year":2010,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Canadian Institutes of Health Research","keywords":"Computer science; User Friendly; Database; Java; Personalization; Interface (matter); Sample (material); Documentation; Source code; World Wide Web; User interface; Operating system","score_opus":0.013319885421856857,"score_gpt":0.2546115274215406,"score_spread":0.24129164199968373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103365019","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028123453,0.000224964,0.15026543,0.0005288326,0.00019657014,0.0006991437,0.020750964,0.8176435,0.0068781893],"genre_scores_gemma":[0.07556572,0.0009098286,0.5662824,0.0045996164,0.00093249243,0.0060199434,0.17678635,0.13326387,0.03563968],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99669516,0.0006331492,0.00038186897,0.00069955154,0.0013252412,0.00026499262],"domain_scores_gemma":[0.9916945,0.0026431275,0.001151566,0.0015511585,0.0018069813,0.0011525922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048386385,0.002319237,0.0016114194,0.003906658,0.0009219825,0.0029849045,0.0051768464,0.0011283669,0.08722075],"category_scores_gemma":[0.012059731,0.0017377125,0.0010433681,0.0018292319,0.00068699033,0.0034196672,0.0051254234,0.0023235315,0.07698797],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014623064,0.00019772224,0.00467504,0.0009594897,0.00015367301,0.00044315332,0.00045193234,0.00165571,0.014263457,0.0027710109,0.7206998,0.2522668],"study_design_scores_gemma":[0.0010609882,0.00035894423,0.010841049,0.00058084814,0.00016521255,0.0010447925,0.00021555526,0.046463225,0.0471887,0.009287062,0.88229835,0.0004953804],"about_ca_topic_score_codex":0.0021217694,"about_ca_topic_score_gemma":0.0016352324,"teacher_disagreement_score":0.08722075,"about_ca_system_score_codex":0.0010414583,"about_ca_system_score_gemma":0.0021344735,"threshold_uncertainty_score":0.2917825},"labels":[],"label_agreement":null},{"id":"W2103765038","doi":"10.1186/1471-2164-14-129","title":"Neurocarta: aggregating and sharing disease-gene relations for the neurosciences","year":2013,"lang":"en","type":"article","venue":"BMC Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of General Medical Sciences; National Institutes of Health","keywords":"Annotation; Context (archaeology); Resource (disambiguation); Inference; Phenotype; Ontology; Biology; Computer science; Computational biology; Data science; Gene; Bioinformatics; Genetics; Artificial intelligence","score_opus":0.032865481258404904,"score_gpt":0.25900990779733557,"score_spread":0.22614442653893066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2103765038","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026988305,0.012297972,0.47707006,0.0060337386,0.0012296025,0.001399745,0.2850646,0.16835615,0.021559833],"genre_scores_gemma":[0.085678324,0.006699456,0.53956443,0.002185591,0.00042965688,0.0017772999,0.35133004,0.008319645,0.004015553],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99601483,0.00092741125,0.0006195095,0.0013491116,0.000889444,0.00019969547],"domain_scores_gemma":[0.98732424,0.00432917,0.0009546052,0.0053326893,0.0012217729,0.0008374505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060608475,0.0021139132,0.0020313815,0.015328132,0.00218049,0.0056592887,0.0030318361,0.0019885208,0.0129535915],"category_scores_gemma":[0.022537798,0.0010938909,0.0042995983,0.010755554,0.0010880459,0.007302943,0.013430434,0.0020883458,0.0069055874],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013443029,0.00035929997,0.032011114,0.006109463,0.0029011327,0.0025227878,0.0037238547,0.01348642,0.013763684,0.05431131,0.45919845,0.41026825],"study_design_scores_gemma":[0.00023375085,0.00020795279,0.017978922,0.0016835078,0.0013888298,0.0018927068,0.0014940606,0.06376896,0.008068447,0.21426168,0.6886668,0.00035432674],"about_ca_topic_score_codex":0.011240883,"about_ca_topic_score_gemma":0.020388264,"teacher_disagreement_score":0.015328132,"about_ca_system_score_codex":0.0018339037,"about_ca_system_score_gemma":0.0048446395,"threshold_uncertainty_score":0.043334067},"labels":[],"label_agreement":null},{"id":"W2104381725","doi":"10.1136/amiajnl-2011-000150","title":"Machine-learned solutions for three stages of clinical information extraction: the state of the art at i2b2 2010","year":2011,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":240,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"U.S. National Library of Medicine","keywords":"Narrative; Benchmark (surveying); Process (computing); Computer science; State (computer science); Data science; Health care; Artificial intelligence; Natural language processing; Political science; Art; Cartography; Literature; Geography; Law","score_opus":0.046539093348759476,"score_gpt":0.34087169057266703,"score_spread":0.29433259722390753,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104381725","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19018315,0.03479063,0.6749485,0.019621221,0.0014087961,0.002710633,0.013085633,0.044424854,0.018826598],"genre_scores_gemma":[0.23016682,0.005052824,0.7174981,0.0020877596,0.0007114905,0.0014008844,0.034645878,0.0012498938,0.0071863723],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9848445,0.00519685,0.0016720366,0.0036379383,0.0039423294,0.0007063025],"domain_scores_gemma":[0.97694415,0.014527619,0.0010612806,0.002330513,0.004370293,0.0007660954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015533911,0.0034274175,0.0018704507,0.004682809,0.0019535907,0.006802565,0.0047907205,0.005742455,0.0048036077],"category_scores_gemma":[0.03633727,0.0010554177,0.0019131855,0.0039096596,0.0011221156,0.006749342,0.0035875633,0.004488295,0.005077991],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009490907,0.0013179533,0.005643842,0.0022189417,0.00036684415,0.00024794755,0.0009187355,0.029674245,0.011834791,0.0025725132,0.054554597,0.8897006],"study_design_scores_gemma":[0.0007935545,0.0013147857,0.01067849,0.000898675,0.00040829094,0.0006474768,0.0016560783,0.83718896,0.048811257,0.022681676,0.07457717,0.00034370366],"about_ca_topic_score_codex":0.011820421,"about_ca_topic_score_gemma":0.013195903,"teacher_disagreement_score":0.015533911,"about_ca_system_score_codex":0.0034478467,"about_ca_system_score_gemma":0.0051366542,"threshold_uncertainty_score":0.08215213},"labels":[],"label_agreement":null},{"id":"W2104990432","doi":"10.1136/amiajnl-2011-000523","title":"The National Center for Biomedical Ontology","year":2011,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":280,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"National Human Genome Research Institute; National Heart, Lung, and Blood Institute; Common Fund; National Institutes of Health","keywords":"Open Biomedical Ontologies; Ontology; Computer science; Biomedicine; Variety (cybernetics); World Wide Web; Data science; Semantic Web; Process ontology; Resource (disambiguation); Upper ontology; Ontology-based data integration; Analytics; Ontology alignment; Bioinformatics; Artificial intelligence","score_opus":0.021415495158009754,"score_gpt":0.3000184411159791,"score_spread":0.2786029459579693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2104990432","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017561116,0.020953612,0.17802796,0.102275245,0.013062319,0.0034147848,0.09079521,0.028219687,0.5614951],"genre_scores_gemma":[0.020236947,0.035738207,0.41962355,0.04490783,0.004230756,0.008361154,0.292733,0.0057479045,0.16842054],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9829101,0.0039202403,0.0027429045,0.0021229954,0.0073265973,0.00097714],"domain_scores_gemma":[0.9584239,0.008738348,0.0026824407,0.011356166,0.013956937,0.004842274],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020001736,0.0017921388,0.0022996562,0.009248925,0.004248622,0.013683201,0.005505059,0.0046647196,0.087827265],"category_scores_gemma":[0.061978642,0.0013212564,0.0019993703,0.010141381,0.0030715128,0.011648498,0.011863334,0.0074638287,0.0842035],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009995461,0.00007831026,0.0009369799,0.0009280579,0.00005921494,0.00015173719,0.00030747126,0.00025307303,0.00050749915,0.15432602,0.6974273,0.14492442],"study_design_scores_gemma":[0.000026624464,0.000012383568,0.00047437422,0.0005185497,0.000020525424,0.000108883,0.000086722124,0.00032626314,0.00015435008,0.032041393,0.9662043,0.000025613605],"about_ca_topic_score_codex":0.016716087,"about_ca_topic_score_gemma":0.010712823,"teacher_disagreement_score":0.087827265,"about_ca_system_score_codex":0.0056849313,"about_ca_system_score_gemma":0.037008505,"threshold_uncertainty_score":0.2938115},"labels":[],"label_agreement":null},{"id":"W2105440823","doi":"10.1093/bfgp/elu015","title":"Event-based text mining for biology and functional genomics","year":2014,"lang":"en","type":"review","venue":"Briefings in Functional Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"Medical Research Council; Wellcome Trust","keywords":"Event (particle physics); Data science; Identification (biology); Computer science; Scalability; Information extraction; Genomics; Function (biology); Biomedical text mining; Biology; Information retrieval; Computational biology; Genome; Data mining; Text mining; Database","score_opus":0.04752817771586031,"score_gpt":0.31421228440090754,"score_spread":0.2666841066850472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105440823","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002151492,0.85297185,0.116663955,0.0072847567,0.0018234443,0.0003803071,0.004851424,0.0021195558,0.011753281],"genre_scores_gemma":[0.015718011,0.82483643,0.13441709,0.003912953,0.0019852014,0.00048771812,0.009439819,0.00025502293,0.008947785],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99884593,0.00026093362,0.0001741302,0.0002476303,0.0004204287,0.0000509945],"domain_scores_gemma":[0.9963198,0.0028750128,0.00026941346,0.00013804325,0.00033780924,0.00006003601],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002221267,0.0012846023,0.0013945338,0.0051833973,0.00039581358,0.0022072196,0.002090949,0.0016800833,0.0054971525],"category_scores_gemma":[0.005124551,0.0004395663,0.0016395068,0.0072507197,0.0009152701,0.0040715295,0.0013632033,0.0020537386,0.0046013705],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004322746,0.000060466256,0.000379441,0.0068313777,0.00016274668,0.0001422022,0.00012253335,0.000959875,0.0022157233,0.008714521,0.031478483,0.9488895],"study_design_scores_gemma":[0.00003986912,0.00006358691,0.0037093204,0.0062190886,0.00024846583,0.0015271587,0.00023816011,0.006072292,0.0056906412,0.055818543,0.92026335,0.00010954607],"about_ca_topic_score_codex":0.0014452086,"about_ca_topic_score_gemma":0.0016008933,"teacher_disagreement_score":0.0054971525,"about_ca_system_score_codex":0.0009159008,"about_ca_system_score_gemma":0.0014796251,"threshold_uncertainty_score":0.018389821},"labels":[],"label_agreement":null},{"id":"W2105612251","doi":"10.1186/1471-2105-12-s8-s12","title":"A linear classifier based on entity recognition tools and a statistical approach to method extraction in the protein-protein interaction literature","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Fundação Luso-Americana para o Desenvolvimento; Natural Sciences and Engineering Research Council of Canada; Queen's University","keywords":"Computer science; Classifier (UML); Artificial intelligence; Pattern recognition (psychology); Relationship extraction; Machine learning; Ranking (information retrieval); Natural language processing; Data mining; Information extraction","score_opus":0.09354681556915123,"score_gpt":0.32892828978095134,"score_spread":0.2353814742118001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105612251","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08532628,0.006431932,0.8770525,0.00242873,0.00054750725,0.00085131516,0.0059535587,0.01530377,0.006104458],"genre_scores_gemma":[0.34361187,0.0019528433,0.63113034,0.00068021147,0.0004916651,0.00092913397,0.015938036,0.00035214578,0.004913682],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99241877,0.0018954229,0.0012920042,0.0017479016,0.0022808718,0.0003650338],"domain_scores_gemma":[0.977759,0.013434252,0.0015045175,0.0015068746,0.0052967654,0.00049857824],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007917657,0.0010038897,0.0010986651,0.013215482,0.0011116564,0.0046059266,0.0020599274,0.0018105032,0.003487555],"category_scores_gemma":[0.023058822,0.00036751124,0.0017524685,0.009767845,0.00087934016,0.0051343082,0.00187229,0.0017673547,0.0052465554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006613724,0.000389346,0.02637365,0.0012143718,0.00029817683,0.0005608738,0.00028600852,0.011517146,0.016001819,0.0071920613,0.024893265,0.9106118],"study_design_scores_gemma":[0.0002204977,0.0011702435,0.024389643,0.0004227853,0.0006184949,0.0019579881,0.00074815156,0.8553806,0.044294897,0.02463037,0.04595951,0.00020678435],"about_ca_topic_score_codex":0.004968896,"about_ca_topic_score_gemma":0.007414908,"teacher_disagreement_score":0.013215482,"about_ca_system_score_codex":0.0016782569,"about_ca_system_score_gemma":0.003477925,"threshold_uncertainty_score":0.041873097},"labels":[],"label_agreement":null},{"id":"W2105731904","doi":"10.1186/2041-1480-4-31","title":"Enhanced XAO: the ontology of Xenopus anatomy and development underpins more accurate annotation of gene expression and queries on Xenbase","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; Division of Emerging Frontiers; National Institute of Child Health and Human Development; National Evolutionary Synthesis Center; National Science Foundation","keywords":"Xenopus; Ontology; Computer science; Annotation; Computational biology; Biology; Information retrieval; Neuroscience; Bioinformatics; Gene; Artificial intelligence; Genetics","score_opus":0.014510146444695544,"score_gpt":0.28589834838558487,"score_spread":0.2713882019408893,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105731904","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020320192,0.0032907566,0.61344033,0.0038825248,0.0012706748,0.0012623635,0.14135467,0.16512777,0.050050575],"genre_scores_gemma":[0.059638932,0.0037680771,0.5445518,0.0025102287,0.00046668723,0.0013375303,0.34244195,0.032844327,0.012440502],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99479926,0.001030547,0.001280803,0.00090103847,0.0016735293,0.0003148683],"domain_scores_gemma":[0.9907286,0.0025305133,0.001016898,0.003494577,0.0018801233,0.0003493166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067498055,0.0016445285,0.0013437567,0.0056354115,0.0013068832,0.0053644306,0.0025143856,0.0015714967,0.018214528],"category_scores_gemma":[0.0155605255,0.0010427922,0.0022200677,0.0051267464,0.001017774,0.009808021,0.0060811327,0.0026102122,0.011801198],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016945126,0.00040399525,0.015464926,0.007922378,0.0005451861,0.0022080364,0.004590542,0.010665162,0.08106614,0.18441851,0.429227,0.2617937],"study_design_scores_gemma":[0.00009005039,0.00006651637,0.004944619,0.0009757167,0.00014343152,0.0007495015,0.0004720578,0.0121316295,0.014671212,0.019215371,0.94637346,0.00016645617],"about_ca_topic_score_codex":0.009939107,"about_ca_topic_score_gemma":0.011286887,"teacher_disagreement_score":0.018214528,"about_ca_system_score_codex":0.0024187579,"about_ca_system_score_gemma":0.004157827,"threshold_uncertainty_score":0.06093365},"labels":[],"label_agreement":null},{"id":"W2105812764","doi":"10.1002/cne.23012","title":"Using text mining to link journal articles to neuroanatomical databases","year":2011,"lang":"en","type":"article","venue":"The Journal of Comparative Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of General Medical Sciences; Canadian Institutes of Health Research","keywords":"Neuroinformatics; Computer science; Lexicon; Information retrieval; Identifier; Standardization; Task (project management); Terminology; Point (geometry); Connectomics; Natural language processing; Data science; Artificial intelligence; Biology; Neuroscience; Linguistics; Connectome","score_opus":0.19802699617624456,"score_gpt":0.36893990552469474,"score_spread":0.1709129093484502,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105812764","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19341673,0.013162357,0.480785,0.0048953965,0.0015630662,0.0046300883,0.24344109,0.026048938,0.032057304],"genre_scores_gemma":[0.17528366,0.006555186,0.66763026,0.0005335316,0.00066116377,0.00223056,0.14121176,0.0018248041,0.0040690475],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99333584,0.0013063833,0.0022567522,0.0012337569,0.0016781134,0.00018920006],"domain_scores_gemma":[0.9243526,0.046661284,0.011560042,0.005671768,0.010753408,0.001000956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009587101,0.0013455073,0.0012522575,0.071572185,0.0019019146,0.005631563,0.0016577061,0.000985572,0.0077623646],"category_scores_gemma":[0.062888615,0.00067195983,0.0013717265,0.058196273,0.000695356,0.005831836,0.0034455177,0.0010204235,0.005143487],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048254343,0.00055003754,0.056009896,0.011877753,0.00090413925,0.002232976,0.003985808,0.00441211,0.03027198,0.013944751,0.074218996,0.8011091],"study_design_scores_gemma":[0.0004041901,0.00059595646,0.15426736,0.003675824,0.0021186755,0.0036392272,0.009373611,0.07959328,0.07245676,0.11407361,0.5592805,0.00052095775],"about_ca_topic_score_codex":0.003501059,"about_ca_topic_score_gemma":0.0045649325,"teacher_disagreement_score":0.071572185,"about_ca_system_score_codex":0.001717325,"about_ca_system_score_gemma":0.0033994303,"threshold_uncertainty_score":0.050702035},"labels":[],"label_agreement":null},{"id":"W2105836514","doi":"10.1186/2041-1480-5-37","title":"CLO: The cell line ontology","year":2014,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":130,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Optech (Canada)","funders":"National Bioscience Database Center; National Institute of Allergy and Infectious Diseases; National Human Genome Research Institute; National Institutes of Health; European Molecular Biology Laboratory; NIH Office of the Director; European Bioinformatics Institute; Japan Science and Technology Agency; University of Michigan","keywords":"Cell culture; Ontology; Cell; Computational biology; Population; Computer science; Biology; Information retrieval; Genetics; Medicine","score_opus":0.00900740289308873,"score_gpt":0.2538138614402925,"score_spread":0.24480645854720376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2105836514","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005271704,0.0036936284,0.6574485,0.013896935,0.0017114504,0.0018386539,0.16575769,0.06628405,0.08409745],"genre_scores_gemma":[0.04337761,0.00640019,0.48271295,0.009105077,0.00080148655,0.0026793503,0.41546544,0.012908039,0.026549794],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956821,0.00079976796,0.00076104933,0.0008120766,0.0015561063,0.0003888695],"domain_scores_gemma":[0.9936212,0.0021649285,0.0005476723,0.0012221741,0.0018074383,0.00063661614],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0048757712,0.0015304777,0.0008701063,0.006064132,0.0021888774,0.0064190514,0.0040130517,0.002914974,0.019125124],"category_scores_gemma":[0.010137587,0.0010866502,0.0022890205,0.00551039,0.001992192,0.012608424,0.0054501174,0.0038750046,0.01733216],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022313034,0.00015350293,0.0022271643,0.0025300863,0.000080520156,0.00072262174,0.0014928109,0.0036162168,0.0054403436,0.37211478,0.46940494,0.141994],"study_design_scores_gemma":[0.0000202429,0.000013253582,0.00044697453,0.0003313004,0.000019071438,0.000314669,0.00017523624,0.0024094058,0.00090794876,0.025201974,0.97011966,0.00004032802],"about_ca_topic_score_codex":0.022007348,"about_ca_topic_score_gemma":0.016553806,"teacher_disagreement_score":0.99358094,"about_ca_system_score_codex":0.004720291,"about_ca_system_score_gemma":0.008050487,"threshold_uncertainty_score":0.06397992},"labels":[],"label_agreement":null},{"id":"W2106350001","doi":"10.1093/bioinformatics/bts350","title":"Identifying aberrant pathways through integrated analysis of knowledge in pharmacogenomics","year":2012,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"National Human Genome Research Institute; National Institutes of Health; European Commission","keywords":"Pharmacogenomics; Ontology; Disease; Drug repositioning; Identification (biology); Computational biology; Computer science; Drug discovery; Data science; Bioinformatics; Drug; Biology; Medicine; Pharmacology","score_opus":0.05877106970867121,"score_gpt":0.3227743530747683,"score_spread":0.26400328336609713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106350001","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12798133,0.005866545,0.8214412,0.0042548226,0.000061234336,0.00045225638,0.024873061,0.005732811,0.009336787],"genre_scores_gemma":[0.40210226,0.0060033016,0.5680935,0.0007082383,0.00006955628,0.00027402476,0.021468807,0.00022994225,0.0010503822],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987435,0.00029446356,0.00017347722,0.00028979382,0.0004360339,0.00006280944],"domain_scores_gemma":[0.99690896,0.0018759187,0.0004366948,0.0003941388,0.00029844965,0.000085820066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020523085,0.00072270655,0.00076368486,0.0071757813,0.0005293324,0.0019803,0.0009143566,0.0005802698,0.0021387374],"category_scores_gemma":[0.0045564556,0.00027283892,0.0012503017,0.008308554,0.0008243576,0.002783734,0.0022308943,0.0008363014,0.00052445685],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014347882,0.000915057,0.08609868,0.0055433707,0.002048925,0.0035278038,0.0011143537,0.08460207,0.07073094,0.078113765,0.012038885,0.65383124],"study_design_scores_gemma":[0.00018675909,0.0002343914,0.04918236,0.0010861051,0.0020010737,0.0027547965,0.0008099175,0.4046531,0.08066529,0.3846085,0.07363238,0.0001852398],"about_ca_topic_score_codex":0.005245926,"about_ca_topic_score_gemma":0.0061256657,"teacher_disagreement_score":0.0071757813,"about_ca_system_score_codex":0.0014261024,"about_ca_system_score_gemma":0.0019741415,"threshold_uncertainty_score":0.010853767},"labels":[],"label_agreement":null},{"id":"W2106537457","doi":"10.1016/j.artmed.2015.03.005","title":"Abstraction networks for terminologies: Supporting management of “big knowledge”","year":2015,"lang":"en","type":"review","venue":"Artificial Intelligence in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":59,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"New York Institute of Technology","funders":"U.S. National Library of Medicine; National Cancer Institute; National Institutes of Health","keywords":"Abstraction; Computer science; Terminology; SNOMED CT; Data science; Variety (cybernetics); Knowledge base; Usability; Artificial intelligence; Information retrieval; Human–computer interaction","score_opus":0.2516891494554278,"score_gpt":0.4751334469324786,"score_spread":0.2234442974770508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106537457","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051799035,0.6668067,0.29670498,0.006171795,0.0010585203,0.00017181749,0.0010603137,0.0009361278,0.02190978],"genre_scores_gemma":[0.04867388,0.6654955,0.27401373,0.0017925629,0.00091531145,0.00024050019,0.003057585,0.00014701937,0.0056638336],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986727,0.00032104726,0.00014424053,0.00019576195,0.0006001506,0.000066034416],"domain_scores_gemma":[0.99643517,0.0022642016,0.0003115456,0.00041471035,0.0004894364,0.00008491116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034913372,0.0007440577,0.0010191964,0.00403288,0.00052001217,0.00305009,0.0025113393,0.001226733,0.0022216316],"category_scores_gemma":[0.0063878605,0.00034409983,0.000852938,0.006347887,0.0017763274,0.0070355274,0.0031172347,0.0017607337,0.0013771771],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026521331,0.000032701475,0.00080400676,0.004251152,0.00011196353,0.000095433425,0.00035747574,0.0016192049,0.0012113425,0.14464748,0.013847469,0.8329952],"study_design_scores_gemma":[0.000013987286,0.00002437919,0.0012007111,0.0042225607,0.00022280359,0.0007503138,0.00032412956,0.0073152054,0.0020080626,0.22921914,0.75464517,0.000053520303],"about_ca_topic_score_codex":0.0020562597,"about_ca_topic_score_gemma":0.0022567473,"teacher_disagreement_score":0.00403288,"about_ca_system_score_codex":0.0013209481,"about_ca_system_score_gemma":0.0029404308,"threshold_uncertainty_score":0.018464208},"labels":[],"label_agreement":null},{"id":"W2106810132","doi":"10.1109/tcbb.2014.2382127","title":"Software Suite for Gene and Protein Annotation Prediction and Similarity Search","year":2014,"lang":"en","type":"article","venue":"IEEE/ACM Transactions on Computational Biology and Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Princess Margaret Cancer Centre; University of Toronto","funders":"","keywords":"Suite; Computer science; Software; Software suite; Annotation; Exploit; Similarity (geometry); Semantic similarity; Web service; Information retrieval; Key (lock); Data mining; Artificial intelligence; World Wide Web; Programming language","score_opus":0.0198074226264382,"score_gpt":0.2824000669206099,"score_spread":0.26259264429417173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2106810132","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0058933794,0.0014936922,0.37396523,0.00052351836,0.0004220737,0.001411248,0.1482622,0.45411882,0.013909835],"genre_scores_gemma":[0.028194426,0.001678637,0.4524836,0.0013337174,0.00014550942,0.006039656,0.4371108,0.05712511,0.0158885],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982495,0.0003042342,0.00026575816,0.00030724265,0.0007235778,0.0001497861],"domain_scores_gemma":[0.99803954,0.000885815,0.00018906523,0.00030460072,0.00039436948,0.00018659918],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021920518,0.0032225072,0.0025907932,0.0033725358,0.0010363952,0.0020520117,0.0038640571,0.0015504366,0.04105811],"category_scores_gemma":[0.005016169,0.0017200665,0.002800505,0.0031961754,0.00069997995,0.0020489194,0.0027926657,0.0029377523,0.03755603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018608244,0.00040884627,0.003894191,0.0041560326,0.0009302907,0.0013234358,0.00036570642,0.0154266525,0.02848469,0.017561896,0.74660206,0.17898534],"study_design_scores_gemma":[0.0019239436,0.000488551,0.0077407854,0.0004487448,0.00042243427,0.0021698824,0.0001970707,0.14574757,0.030910116,0.06501504,0.7444863,0.00044964397],"about_ca_topic_score_codex":0.0051962356,"about_ca_topic_score_gemma":0.005769588,"teacher_disagreement_score":0.04105811,"about_ca_system_score_codex":0.0010994106,"about_ca_system_score_gemma":0.002957284,"threshold_uncertainty_score":0.137353},"labels":[],"label_agreement":null},{"id":"W2107141268","doi":"10.1093/bioinformatics/btn381","title":"Multi-dimensional classification of biomedical text: Toward automated, practical provision of high-utility text to diverse users","year":2008,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":134,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; National Science Foundation","keywords":"Computer science; Categorization; Annotation; Sentence; Information retrieval; Task (project management); Variety (cybernetics); Artificial intelligence; Information extraction; Natural language processing; Focus (optics)","score_opus":0.05631967128910629,"score_gpt":0.3222240203340662,"score_spread":0.2659043490449599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107141268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07253602,0.0010016214,0.90424407,0.0036080026,0.00013496754,0.001111048,0.0030093428,0.010885166,0.003469885],"genre_scores_gemma":[0.14472787,0.00037584556,0.8456869,0.0005182525,0.0002553826,0.00066216127,0.0057548834,0.0004666181,0.0015521537],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99162674,0.0034674087,0.0009959948,0.0018059567,0.0018306477,0.0002732264],"domain_scores_gemma":[0.9609073,0.021452108,0.004210506,0.0042091208,0.008067651,0.0011533013],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008546948,0.0014343419,0.0014012114,0.008093144,0.0017650661,0.004641315,0.0027960713,0.002007478,0.002847018],"category_scores_gemma":[0.032611687,0.0005975825,0.0010046676,0.004910848,0.0017556103,0.007196731,0.0043941927,0.0024753155,0.003803959],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072279293,0.00060465804,0.012592132,0.0012587814,0.00007070352,0.00041029786,0.0043614176,0.0060200817,0.05628081,0.008541122,0.020573912,0.8885633],"study_design_scores_gemma":[0.00018279371,0.0007082168,0.020288693,0.00061044656,0.00020832759,0.0015258769,0.0074936696,0.74549764,0.08810277,0.070644416,0.06445085,0.0002863524],"about_ca_topic_score_codex":0.0020124393,"about_ca_topic_score_gemma":0.0022142753,"teacher_disagreement_score":0.008546948,"about_ca_system_score_codex":0.0014648185,"about_ca_system_score_gemma":0.0019612666,"threshold_uncertainty_score":0.045201123},"labels":[],"label_agreement":null},{"id":"W2107658719","doi":"10.1109/hicss.2011.61","title":"An Ontology-Based Electronic Medical Record for Chronic Disease Management","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University","funders":"Dalhousie University","keywords":"Ontology; Interoperability; Computer science; Semantic interoperability; Clinical decision support system; Ontology-based data integration; Disease; Upper ontology; Knowledge management; Decision support system; Chronic disease; Medical record; Semantic Web; Information retrieval; World Wide Web; Medicine; Data mining; Family medicine; Pathology","score_opus":0.021379572446407077,"score_gpt":0.292642725994336,"score_spread":0.2712631535479289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2107658719","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0074484493,0.0018093158,0.9177378,0.009696038,0.000703923,0.0012462124,0.009463023,0.002502255,0.049393024],"genre_scores_gemma":[0.051251363,0.0021927017,0.92537135,0.0013639351,0.00018823064,0.000549925,0.011191259,0.00018525585,0.0077059497],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9961212,0.00107594,0.00083680655,0.00044601853,0.0013292701,0.00019077677],"domain_scores_gemma":[0.99659616,0.00096464745,0.00043073273,0.00092948566,0.0008572019,0.00022175995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004431562,0.00035472692,0.0005508376,0.003378757,0.0019862195,0.0047968263,0.0014816005,0.0015099583,0.0034872948],"category_scores_gemma":[0.007517672,0.00033678822,0.001530291,0.0051690917,0.0013379984,0.006110482,0.00254195,0.0020232073,0.001623113],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000737983,0.00017301674,0.002891014,0.00094416563,0.00011603527,0.000743642,0.0019095901,0.004770453,0.006426317,0.70999646,0.053535614,0.21841994],"study_design_scores_gemma":[0.000048365415,0.000051289575,0.00322531,0.00077619555,0.00017897176,0.0011815191,0.0007959624,0.026974132,0.003823003,0.14160839,0.8212295,0.00010729211],"about_ca_topic_score_codex":0.0156154055,"about_ca_topic_score_gemma":0.016840698,"teacher_disagreement_score":0.0156154055,"about_ca_system_score_codex":0.0025266297,"about_ca_system_score_gemma":0.008383489,"threshold_uncertainty_score":0.031049013},"labels":[],"label_agreement":null},{"id":"W2108317086","doi":"10.2147/ijn.s4375","title":"Spatiotemporal integration of molecular and anatomical data in virtual reality using semantic mapping","year":2009,"lang":"en","type":"article","venue":"International Journal of Nanomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Government of Canada; Genome Alberta; Government of Alberta; Genome Canada","keywords":"Computer science; Ontology; Inference; Semantic reasoner; Data integration; Context (archaeology); Semantic mapping; Semantic integration; Visualization; Virtual reality; Data science; Information retrieval; Artificial intelligence; Data mining; Semantic computing; Semantic Web; Biology","score_opus":0.04994161222320141,"score_gpt":0.3538430991709638,"score_spread":0.30390148694776237,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108317086","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006500046,0.00007427509,0.9912952,0.00023102788,0.000019710524,0.0000547958,0.00023811201,0.0005184643,0.0010682554],"genre_scores_gemma":[0.16984336,0.00029293765,0.82837456,0.00007973953,0.000015113177,0.0001830416,0.0006100089,0.000113094015,0.00048817164],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981596,0.0006872976,0.0001923973,0.00030483818,0.00057147845,0.0000843652],"domain_scores_gemma":[0.9978144,0.0011159756,0.00022400242,0.00046531472,0.0002745192,0.000105814775],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002756563,0.00057439297,0.00076361315,0.003189609,0.00072490936,0.0045112767,0.0013898172,0.0007224036,0.0019724944],"category_scores_gemma":[0.0076384824,0.0005756107,0.0024717364,0.0023384932,0.0017809574,0.0034805636,0.004165449,0.0011162609,0.00033373313],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023255167,0.00017765803,0.004886286,0.00062516757,0.00036258207,0.0010961349,0.0036261496,0.4454126,0.02052757,0.30195898,0.0033084564,0.21778591],"study_design_scores_gemma":[0.00004430114,0.00011285365,0.0019140585,0.00015304059,0.00012968385,0.0005744473,0.0010409842,0.7860865,0.012343028,0.16854206,0.028928207,0.00013082442],"about_ca_topic_score_codex":0.0063639595,"about_ca_topic_score_gemma":0.006460991,"teacher_disagreement_score":0.0063639595,"about_ca_system_score_codex":0.00094084797,"about_ca_system_score_gemma":0.001523108,"threshold_uncertainty_score":0.014578283},"labels":[],"label_agreement":null},{"id":"W2108816030","doi":"10.1136/ebm.10.4.101","title":"Finding the gold in Medline: clinical queries","year":2005,"lang":"en","type":"editorial","venue":"Evidence-Based Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"MEDLINE; Computer science; Treasure; Search engine indexing; Information retrieval; World Wide Web; Data science; History; Political science","score_opus":0.07816964072642624,"score_gpt":0.4094760960115682,"score_spread":0.33130645528514197,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2108816030","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00017610766,0.03380458,0.0015138614,0.23488167,0.7255152,0.00015530661,0.00046502115,0.00037625892,0.003112077],"genre_scores_gemma":[0.0024776675,0.06003497,0.0033304654,0.17024161,0.74336237,0.00027236686,0.0006411071,0.0003757677,0.01926363],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9916328,0.0025761595,0.0022283783,0.00067642983,0.002709877,0.00017622771],"domain_scores_gemma":[0.9441994,0.0361426,0.0023666571,0.0008211144,0.013452483,0.0030177354],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015483951,0.0024312288,0.0030789857,0.011359791,0.0023559995,0.008707066,0.002906923,0.009206734,0.013546886],"category_scores_gemma":[0.07179484,0.0012830996,0.0020796699,0.0061332304,0.0023737757,0.006209508,0.0020501553,0.012360524,0.008187511],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002123914,0.00000633613,0.000021954826,0.0004918459,0.000019313871,0.00007105722,0.000038154438,0.000023398978,0.000044309843,0.00024541444,0.99131167,0.007705344],"study_design_scores_gemma":[0.00007015897,0.000027741622,0.0002580161,0.002058989,0.000102581624,0.00031184452,0.0001323601,0.00019854831,0.00014215776,0.0012554146,0.995411,0.00003129452],"about_ca_topic_score_codex":0.0034528442,"about_ca_topic_score_gemma":0.009099022,"teacher_disagreement_score":0.984516,"about_ca_system_score_codex":0.005438401,"about_ca_system_score_gemma":0.006471211,"threshold_uncertainty_score":0.0818879},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"not_applicable","genre":"editorial","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"not_applicable","genre":"editorial","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"agree"},{"id":"W2109429447","doi":"10.1186/1471-2105-12-486","title":"Constructing a semantic predication gold standard from the biomedical literature","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Unified Medical Language System; Computer science; Information retrieval; Annotation; Terminology; Natural language processing; Biomedical text mining; Ontology; Task (project management); Semantic similarity; Controlled vocabulary; Artificial intelligence; Text mining; Linguistics","score_opus":0.023971190296094843,"score_gpt":0.24545670248368603,"score_spread":0.2214855121875912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109429447","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39241973,0.007573518,0.5524862,0.0029189938,0.00097527495,0.0026049728,0.010811068,0.004138265,0.026071975],"genre_scores_gemma":[0.55591816,0.0012768056,0.41516066,0.0005459473,0.00023748311,0.0036922528,0.020997494,0.0008516093,0.0013195588],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9089855,0.03444411,0.018497694,0.013586818,0.02311136,0.0013745527],"domain_scores_gemma":[0.62771255,0.22673105,0.019229854,0.036286756,0.088023774,0.0020160328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10395984,0.0013044638,0.0016064391,0.032314185,0.005148577,0.005814632,0.0031005915,0.0026264763,0.003257273],"category_scores_gemma":[0.25596792,0.0006486081,0.002040533,0.01437126,0.003850662,0.0066217887,0.0100791715,0.0021021303,0.0020031712],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016499958,0.0008509285,0.15301399,0.016171446,0.0021491863,0.001959697,0.031692017,0.013988637,0.07400356,0.057698056,0.03459436,0.6122282],"study_design_scores_gemma":[0.0004357977,0.001015791,0.26123166,0.008776964,0.002800792,0.002983322,0.023733873,0.15527298,0.18070926,0.14711143,0.21511935,0.00080886047],"about_ca_topic_score_codex":0.0028117222,"about_ca_topic_score_gemma":0.004137385,"teacher_disagreement_score":0.10395984,"about_ca_system_score_codex":0.0032495542,"about_ca_system_score_gemma":0.006815845,"threshold_uncertainty_score":0.5497988},"labels":[],"label_agreement":null},{"id":"W2109977913","doi":"10.1197/jamia.m3083","title":"Description of a Rule-based System for the i2b2 Challenge in Natural Language Processing for Clinical Data","year":2009,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Computer science; Informatics; Health informatics; Domain (mathematical analysis); Rule-based system; Replicate; Natural language processing; Artificial intelligence; Data science; Medicine; Pathology; Engineering; Public health","score_opus":0.048107105421720676,"score_gpt":0.3755095698646421,"score_spread":0.3274024644429214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2109977913","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005444541,0.00017030741,0.9064571,0.0018201293,0.00023481787,0.0024516166,0.0019608312,0.07821562,0.0032450065],"genre_scores_gemma":[0.039562017,0.00012834794,0.9471949,0.0015282881,0.00012185858,0.0013731617,0.0040093227,0.002192108,0.0038898883],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99316496,0.0015504584,0.0013051202,0.0016808765,0.002084681,0.00021391662],"domain_scores_gemma":[0.98325306,0.009571032,0.0005641478,0.0024927321,0.0031884003,0.00093053235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009983857,0.0014548696,0.0017347555,0.0022241073,0.0015925962,0.0066173235,0.005750418,0.0045603714,0.018051999],"category_scores_gemma":[0.0267435,0.0013634339,0.001646734,0.0015479788,0.0011345175,0.0045759836,0.0026783491,0.003918643,0.013545759],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014868148,0.0022240735,0.010545838,0.0016510945,0.00086067687,0.006245217,0.0017353384,0.0515382,0.070541516,0.022071071,0.18183205,0.64926803],"study_design_scores_gemma":[0.00071796187,0.0004340702,0.0019834645,0.00029608558,0.00021199345,0.0022272614,0.0003556751,0.8100956,0.047800522,0.023125378,0.11241495,0.00033696246],"about_ca_topic_score_codex":0.007554048,"about_ca_topic_score_gemma":0.0071433336,"teacher_disagreement_score":0.018051999,"about_ca_system_score_codex":0.001165716,"about_ca_system_score_gemma":0.0031367438,"threshold_uncertainty_score":0.060389996},"labels":[],"label_agreement":null},{"id":"W2110217168","doi":"10.1109/e-science.2009.29","title":"Comparing METS and OAI-ORE for Encapsulating Scientific Data Products: A Protein Crystallography Case Study","year":2009,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Cambridge; University of Queensland; University of Southampton; Pennsylvania State University; Ryerson University","keywords":"Workflow; Computer science; Pipeline (software); Protein Data Bank; Information retrieval; Data science; World Wide Web; Database; Chemistry; Protein structure; Programming language","score_opus":0.0935972833186771,"score_gpt":0.3360682235167509,"score_spread":0.2424709401980738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110217168","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.57436174,0.0018226587,0.38283113,0.0047856076,0.00021906316,0.0015286381,0.0013299866,0.004876506,0.028244674],"genre_scores_gemma":[0.6252184,0.0011781516,0.3659049,0.0003704195,0.000046149486,0.00038375417,0.0020219274,0.0007180776,0.004158297],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9871652,0.005223729,0.0014568255,0.0004968101,0.0050950455,0.0005624093],"domain_scores_gemma":[0.97002745,0.018919364,0.0018539302,0.0052185594,0.0034059687,0.0005747989],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.018805297,0.00059282803,0.000637126,0.002346202,0.0011641474,0.00513285,0.0017613608,0.0017089979,0.0010704405],"category_scores_gemma":[0.034741774,0.00047736286,0.0011360159,0.003689779,0.0018936592,0.008774388,0.0040410263,0.0016970162,0.0004519525],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042110467,0.0023366096,0.08213757,0.0035874462,0.00042164297,0.0057763183,0.013155613,0.0743171,0.034320373,0.30285475,0.012985882,0.4638956],"study_design_scores_gemma":[0.0006252848,0.0027256312,0.031779137,0.0011042529,0.00057272695,0.0068686516,0.02628654,0.550356,0.089395456,0.09265037,0.19723439,0.00040168763],"about_ca_topic_score_codex":0.0066095763,"about_ca_topic_score_gemma":0.00820837,"teacher_disagreement_score":0.99486715,"about_ca_system_score_codex":0.0019485154,"about_ca_system_score_gemma":0.0024120696,"threshold_uncertainty_score":0.09945309},"labels":[],"label_agreement":null},{"id":"W2110279208","doi":"10.1371/journal.pone.0025513","title":"The Chemical Information Ontology: Provenance and Disambiguation for Chemical Data on the Biological Semantic Web","year":2011,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":136,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Biotechnology and Biological Sciences Research Council; Uppsala Universitet","keywords":"Computer science; Cheminformatics; Ontology; Context (archaeology); Semantic Web; Open Biomedical Ontologies; Data science; Domain (mathematical analysis); Ontology-based data integration; Web Ontology Language; Information retrieval; World Wide Web; Suggested Upper Merged Ontology; Bioinformatics","score_opus":0.13324293468670073,"score_gpt":0.26382207149076714,"score_spread":0.1305791368040664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110279208","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030818998,0.00050950574,0.9709707,0.0035251847,0.00046526067,0.00060809695,0.005786178,0.0065408824,0.008512382],"genre_scores_gemma":[0.048443165,0.0019608259,0.92417777,0.0011069648,0.00026949187,0.00050989504,0.015386166,0.0020812477,0.0060644606],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920357,0.0017699016,0.0012305275,0.00095048174,0.0035931976,0.00042030128],"domain_scores_gemma":[0.9870918,0.0035519965,0.0010601563,0.0045564044,0.0029818232,0.0007578221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012477131,0.00076451636,0.0011579914,0.00793828,0.0034967677,0.008156353,0.0027270375,0.0025238935,0.003390196],"category_scores_gemma":[0.025463402,0.0010429557,0.002197551,0.008294956,0.003471192,0.021466741,0.006756967,0.0048307623,0.0024114512],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020623313,0.00017791535,0.0027721866,0.00069318875,0.00008007886,0.0011349488,0.0021084123,0.011578394,0.0037988303,0.7286191,0.046608448,0.20222218],"study_design_scores_gemma":[0.00006565531,0.000032257205,0.0011940404,0.00068744685,0.000087912085,0.0007921446,0.000745661,0.079009585,0.0075878394,0.38919902,0.5204747,0.0001237234],"about_ca_topic_score_codex":0.021595258,"about_ca_topic_score_gemma":0.018091315,"teacher_disagreement_score":0.021595258,"about_ca_system_score_codex":0.0042108586,"about_ca_system_score_gemma":0.009712764,"threshold_uncertainty_score":0.06598616},"labels":[],"label_agreement":null},{"id":"W2110301462","doi":"10.1093/bib/bbs053","title":"Evaluation of research in biomedical ontologies","year":2012,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":81,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Association of Occupational Therapists; Carleton University","funders":"National Human Genome Research Institute","keywords":"Computer science; Ontology; Open Biomedical Ontologies; Terminology; Biomedicine; IDEF5; Data science; Consistency (knowledge bases); Controlled vocabulary; Domain (mathematical analysis); Information retrieval; Semantic Web; Upper ontology; Ontology alignment; Artificial intelligence; Bioinformatics","score_opus":0.14574975016251152,"score_gpt":0.42246491525679264,"score_spread":0.2767151650942811,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110301462","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55464745,0.054113623,0.1926326,0.021129217,0.0018882753,0.012933898,0.010628298,0.00091185025,0.1511147],"genre_scores_gemma":[0.8184204,0.006801391,0.16116463,0.0010667368,0.0003659591,0.0055731577,0.0037172034,0.00019461787,0.002695877],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.6204483,0.23919094,0.044816915,0.008143893,0.08410408,0.003295944],"domain_scores_gemma":[0.3892099,0.42606604,0.041348826,0.019538775,0.119483374,0.0043530776],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.24989294,0.0019188859,0.0026854326,0.02983064,0.0023605342,0.013171856,0.0022099633,0.0032715236,0.0040571443],"category_scores_gemma":[0.48472667,0.0005589587,0.0033812516,0.024335317,0.0041941307,0.009072207,0.006845463,0.0013164737,0.00072065566],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0069838157,0.001582761,0.14585453,0.022922777,0.008562799,0.00078649895,0.008850274,0.025994912,0.0068019647,0.11318254,0.011715298,0.6467619],"study_design_scores_gemma":[0.0027848803,0.011024934,0.2527389,0.020371113,0.023247384,0.001499499,0.03400346,0.111359194,0.0681212,0.28665137,0.18717007,0.0010280408],"about_ca_topic_score_codex":0.0033812886,"about_ca_topic_score_gemma":0.0039362037,"teacher_disagreement_score":0.75010705,"about_ca_system_score_codex":0.013034875,"about_ca_system_score_gemma":0.010541995,"threshold_uncertainty_score":0.92501557},"labels":[],"label_agreement":null},{"id":"W2110403719","doi":"10.1109/cic.2005.1588195","title":"Diagnostic imaging on demand, information management insights based on canadian DI/EHR strategies","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Stakeholder; Electronic health record; Process (computing); Process management; Business; Healthcare delivery; Stakeholder engagement; Knowledge management; Health care; On demand; Maturity (psychological); Information management; Computer science; Public relations; Psychology; Political science","score_opus":0.0050542910146237465,"score_gpt":0.22371491284414366,"score_spread":0.21866062182951992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110403719","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2460055,0.001126514,0.03458981,0.072411865,0.000083769366,0.00046355088,0.003059458,0.00048499866,0.6417746],"genre_scores_gemma":[0.9353033,0.0013120916,0.03346244,0.0020273493,0.00002614098,0.00009271611,0.0017517084,0.000081119775,0.025943214],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99695,0.0006928005,0.0001828171,0.00017116392,0.001368992,0.00063434715],"domain_scores_gemma":[0.9946877,0.0014463641,0.00022745227,0.0003240388,0.0026285583,0.00068592146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004883269,0.0002840339,0.00014066348,0.0033231608,0.0026403868,0.007239491,0.0011482768,0.00081007194,0.0055551883],"category_scores_gemma":[0.00988664,0.00019337682,0.0002515561,0.004469171,0.0014206722,0.0031845097,0.0024421108,0.00091405486,0.00058338745],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001806299,0.00026064375,0.04728304,0.00022391407,0.0000230644,0.0010024869,0.015835106,0.00636392,0.0033236519,0.65600127,0.056861423,0.2126409],"study_design_scores_gemma":[0.00008544204,0.00012468707,0.08555794,0.00037266043,0.000086331806,0.00076317764,0.0761075,0.067376256,0.008728535,0.16801496,0.5925628,0.00021973027],"about_ca_topic_score_codex":0.7420924,"about_ca_topic_score_gemma":0.8039093,"teacher_disagreement_score":0.9599825,"about_ca_system_score_codex":0.04001749,"about_ca_system_score_gemma":0.0370642,"threshold_uncertainty_score":0.51885295},"labels":[],"label_agreement":null},{"id":"W2110961693","doi":"10.1136/amiajnl-2014-002901","title":"Using the wisdom of the crowds to find critical errors in biomedical ontologies: a study of SNOMED CT","year":2014,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":62,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"U.S. National Library of Medicine; National Institute of General Medical Sciences; National Human Genome Research Institute; Common Fund; National Institutes of Health","keywords":"SNOMED CT; Computer science; Systematized Nomenclature of Medicine; Ontology; Scalability; Domain (mathematical analysis); Crowdsourcing; Artificial intelligence; Open Biomedical Ontologies; Subject-matter expert; Information retrieval; Machine learning; Data science; Terminology; Domain knowledge; Ontology alignment; Process ontology; World Wide Web; Expert system; Database","score_opus":0.021010209747288407,"score_gpt":0.340336319343528,"score_spread":0.3193261095962396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2110961693","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9133872,0.0015048691,0.06693554,0.005299812,0.0002495641,0.0016536252,0.0005812004,0.00034674475,0.010041568],"genre_scores_gemma":[0.94947386,0.0004181383,0.04559205,0.0019297946,0.0001399828,0.00056693144,0.00042789907,0.00019146287,0.0012599097],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.90499276,0.06721622,0.0032789917,0.008587825,0.014729924,0.0011943214],"domain_scores_gemma":[0.45065904,0.464858,0.025863836,0.026947875,0.027658578,0.0040126718],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08405937,0.001241073,0.0010616478,0.0064628315,0.0053254194,0.003678222,0.0040859347,0.0035130905,0.0022055528],"category_scores_gemma":[0.3054589,0.0009613444,0.0014160728,0.0030209636,0.0073763235,0.007817146,0.0068185506,0.0033668221,0.0006922477],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032627368,0.0042827767,0.385748,0.002932348,0.0014414842,0.00458063,0.20210195,0.026158936,0.006497864,0.023570621,0.020903533,0.31851915],"study_design_scores_gemma":[0.0013221033,0.0041638757,0.25618505,0.0033967097,0.0010917436,0.0054295887,0.16370098,0.32649183,0.016649766,0.12954523,0.090792954,0.0012301375],"about_ca_topic_score_codex":0.024688445,"about_ca_topic_score_gemma":0.01989474,"teacher_disagreement_score":0.08405937,"about_ca_system_score_codex":0.004686967,"about_ca_system_score_gemma":0.0055779913,"threshold_uncertainty_score":0.4445538},"labels":[],"label_agreement":null},{"id":"W2111255192","doi":"","title":"Applying Probabilistic Thematic Clustering for Classification in the TREC 2005 Genomics Track","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Computer science; Artificial intelligence; Feature selection; Classifier (UML); Cluster analysis; Categorization; Naive Bayes classifier; Test set; Machine learning; Pattern recognition (psychology); Data mining; Natural language processing; Support vector machine","score_opus":0.051307684968168746,"score_gpt":0.302721482534822,"score_spread":0.2514137975666533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2111255192","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36450386,0.0019736225,0.6024582,0.0023698627,0.00046146073,0.001494879,0.0056749303,0.010488784,0.010574351],"genre_scores_gemma":[0.5215672,0.00041801023,0.46528304,0.0002941605,0.00014760178,0.0005349255,0.0084757665,0.00029567519,0.0029836944],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9916414,0.004200011,0.0004941955,0.0011031441,0.0020299852,0.0005312413],"domain_scores_gemma":[0.9901183,0.005988348,0.00039488636,0.0007941818,0.0024638001,0.00024044467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009641312,0.00087687536,0.00097509346,0.0047490704,0.0019804074,0.0018039485,0.0017390688,0.0017086586,0.0013156991],"category_scores_gemma":[0.016314887,0.00037015715,0.0014042064,0.003975792,0.0006454828,0.0026499752,0.0009505239,0.0012758456,0.0010322522],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012631309,0.00092438987,0.027168397,0.0005539322,0.00034368955,0.00041070717,0.0014107474,0.089249656,0.022653496,0.0084930565,0.04778267,0.79974616],"study_design_scores_gemma":[0.00012260918,0.00034870597,0.01781772,0.0000661966,0.00013951625,0.00035302612,0.0012381485,0.9255542,0.024211215,0.014969652,0.015067454,0.00011160566],"about_ca_topic_score_codex":0.027753886,"about_ca_topic_score_gemma":0.025311263,"teacher_disagreement_score":0.027753886,"about_ca_system_score_codex":0.002594422,"about_ca_system_score_gemma":0.0020994223,"threshold_uncertainty_score":0.055184662},"labels":[],"label_agreement":null},{"id":"W2112455505","doi":"10.5339/qfarf.2013.biop-034","title":"The Advice Infrastructure For Generating And Delivering Evidence-Informed Clinical Decision Support Services: A Knowledge Management Approach","year":2013,"lang":"en","type":"article","venue":"Qatar Foundation Annual Research Forum Volume 2013 Issue 1","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Clinical decision support system; Decision support system; Knowledge management; Health care; Analytics; Operationalization; Decision aids; Computer science; Process management; Medicine; Business; Data science; Artificial intelligence","score_opus":0.04901319358007845,"score_gpt":0.41971115863003833,"score_spread":0.3706979650499599,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2112455505","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038904129,0.00055532774,0.96345055,0.010873408,0.000104628634,0.0006442735,0.0009064392,0.0070374156,0.012537619],"genre_scores_gemma":[0.063227,0.0010610578,0.9269627,0.0016517291,0.00016276022,0.0005417521,0.0028127488,0.0004441814,0.0031360814],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9906887,0.0034656818,0.0013473022,0.0015146652,0.0023370786,0.00064663193],"domain_scores_gemma":[0.9846058,0.0076371976,0.0011636488,0.0030015362,0.0024944183,0.001097418],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012959097,0.0010470583,0.0008312941,0.005017097,0.0019285638,0.011012424,0.005066751,0.0038103326,0.005882658],"category_scores_gemma":[0.022743223,0.0012704367,0.0018167702,0.004629597,0.0030460937,0.00956912,0.0072587878,0.0038994816,0.0031557833],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021785109,0.00051014614,0.003967698,0.001593303,0.00026182464,0.0009039103,0.003151132,0.041934714,0.007611423,0.501985,0.040804606,0.3970584],"study_design_scores_gemma":[0.00015282659,0.00014778301,0.0017439085,0.0011246224,0.00018116279,0.0005048109,0.0013254671,0.18728347,0.010916599,0.38921487,0.40724936,0.00015514246],"about_ca_topic_score_codex":0.0101985745,"about_ca_topic_score_gemma":0.006700467,"teacher_disagreement_score":0.012959097,"about_ca_system_score_codex":0.004027655,"about_ca_system_score_gemma":0.013924754,"threshold_uncertainty_score":0.06853503},"labels":[],"label_agreement":null},{"id":"W2114975391","doi":"10.4056/sigs.2025347","title":"Minimal Information for Neural Electromagnetic Ontologies (MINEMO): A standards-compliant method for analysis and integration of event-related potentials (ERP) data","year":2011,"lang":"en","type":"article","venue":"Standards in Genomic Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"National Institute of Biomedical Imaging and Bioengineering","keywords":"Metadata; Computer science; Upload; Ontology; RDF; Context (archaeology); Event (particle physics); Information retrieval; Semantic integration; Semantic Web; World Wide Web","score_opus":0.058946493492542634,"score_gpt":0.3771644178397519,"score_spread":0.3182179243472093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2114975391","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020281228,0.0001761222,0.98505473,0.0005292135,0.000058640835,0.0006711758,0.004415182,0.0045785443,0.0024882767],"genre_scores_gemma":[0.013098228,0.00019563588,0.9747252,0.00021457831,0.000035977675,0.0012687558,0.008722097,0.00090078276,0.0008386975],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9807451,0.0053674313,0.0054549677,0.0019582217,0.0060007162,0.0004735389],"domain_scores_gemma":[0.9666463,0.013642925,0.0041096644,0.009371325,0.00553116,0.00069854496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018698527,0.0017655285,0.0016829459,0.013256124,0.0031427925,0.0070384424,0.0036074694,0.0026798258,0.0042845467],"category_scores_gemma":[0.058472455,0.0018940683,0.003998527,0.008673493,0.002018107,0.011543689,0.009948552,0.0044914344,0.0023709293],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037779045,0.0004169888,0.0056316224,0.0034203595,0.0005686251,0.0014394902,0.006247142,0.0058682575,0.016079027,0.40673003,0.053542845,0.49967775],"study_design_scores_gemma":[0.00012798741,0.00012944681,0.0036678489,0.0020781294,0.0003791592,0.002143911,0.0016439622,0.04901135,0.027103683,0.42085776,0.49252015,0.00033661426],"about_ca_topic_score_codex":0.003679285,"about_ca_topic_score_gemma":0.0064550284,"teacher_disagreement_score":0.018698527,"about_ca_system_score_codex":0.0024124975,"about_ca_system_score_gemma":0.008849358,"threshold_uncertainty_score":0.09888846},"labels":[],"label_agreement":null},{"id":"W2115200041","doi":"10.1504/ijbra.2007.015009","title":"Enhanced semantic access to the protein engineering literature using ontologies populated by text mining","year":2007,"lang":"en","type":"article","venue":"International Journal of Bioinformatics Research and Applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Concordia University","keywords":"Computer science; Ontology; Workflow; Information retrieval; Open Biomedical Ontologies; World Wide Web; Natural language processing; Semantic Web; Data science; Upper ontology; Suggested Upper Merged Ontology; Database","score_opus":0.04049951468218915,"score_gpt":0.39282512130567393,"score_spread":0.3523256066234848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115200041","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13738915,0.00090563716,0.8214025,0.003007169,0.00006763369,0.00052390347,0.008027791,0.018181052,0.010495151],"genre_scores_gemma":[0.27417555,0.0010320783,0.7047298,0.00053023855,0.00007049192,0.00022800404,0.016521245,0.0006176118,0.0020950283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978933,0.0007138605,0.00034342264,0.00028620928,0.0006494371,0.00011384249],"domain_scores_gemma":[0.9896095,0.0063138995,0.00077755254,0.0016049782,0.0013719351,0.00032217085],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.003475863,0.0006198809,0.0008512216,0.007875004,0.0009931797,0.004200508,0.0011369523,0.0009195677,0.0018143929],"category_scores_gemma":[0.013638683,0.00043009932,0.0010082446,0.008066084,0.00074530026,0.006961498,0.0041130846,0.0008288011,0.0009260442],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012384374,0.0010501278,0.016547842,0.0027985496,0.00029002692,0.0040630256,0.008517596,0.020490903,0.10746151,0.08640626,0.021833822,0.72930187],"study_design_scores_gemma":[0.0003566046,0.00029797718,0.010398673,0.0007669788,0.0006646114,0.004298457,0.0066180257,0.43325344,0.12943769,0.1889948,0.22457013,0.0003426344],"about_ca_topic_score_codex":0.0026991565,"about_ca_topic_score_gemma":0.004416994,"teacher_disagreement_score":0.9957995,"about_ca_system_score_codex":0.0007003162,"about_ca_system_score_gemma":0.0018754972,"threshold_uncertainty_score":0.01838231},"labels":[],"label_agreement":null},{"id":"W2115504560","doi":"10.1186/gb-2008-9-2-r31","title":"Text-mining assisted regulatory annotation","year":2008,"lang":"en","type":"article","venue":"Genome biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"Natural Sciences and Engineering Research Council of Canada; Vlaamse regering; Fonds Wetenschappelijk Onderzoek; Genome British Columbia; Canadian Institutes of Health Research; Genome Canada; Michael Smith Health Research BC; National Evolutionary Synthesis Center; European Molecular Biology Organization; National Science Foundation","keywords":"Annotation; Regulatory sequence; Ranking (information retrieval); Computational biology; Computer science; Genome; Gene Annotation; Relevance (law); Genomics; Information retrieval; Gene; Data mining; Biology; Regulation of gene expression; Genetics; Artificial intelligence","score_opus":0.030476990184004252,"score_gpt":0.27026598582879635,"score_spread":0.2397889956447921,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2115504560","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08284151,0.00527589,0.820248,0.003587005,0.00050892064,0.0017249794,0.052545916,0.022656795,0.010611041],"genre_scores_gemma":[0.1912042,0.0025822762,0.73622364,0.0005152858,0.00035223487,0.0012781086,0.06329605,0.0005758854,0.0039722077],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9975169,0.00067296013,0.00041664514,0.00066516496,0.0006350318,0.00009340823],"domain_scores_gemma":[0.98599374,0.008835962,0.0014485588,0.0007745904,0.0027447059,0.00020247004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029759025,0.0013759263,0.0009988334,0.008152766,0.00085759046,0.0017359029,0.0017834014,0.0010624187,0.0059192167],"category_scores_gemma":[0.013561524,0.0002464964,0.0012247255,0.006322403,0.00055748963,0.0020568478,0.0010026959,0.0010398711,0.002522241],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007865641,0.0006542987,0.010866372,0.0075720577,0.00035977972,0.0015409249,0.0009716527,0.01958794,0.08249043,0.01428594,0.04174977,0.81913424],"study_design_scores_gemma":[0.00038234,0.00042884087,0.019693127,0.001349882,0.0009294017,0.0028942577,0.0012786131,0.47822076,0.21314184,0.08328139,0.19816013,0.00023936771],"about_ca_topic_score_codex":0.0018576,"about_ca_topic_score_gemma":0.0019903309,"teacher_disagreement_score":0.008152766,"about_ca_system_score_codex":0.0009436408,"about_ca_system_score_gemma":0.002200536,"threshold_uncertainty_score":0.019801736},"labels":[],"label_agreement":null},{"id":"W2116420111","doi":"10.1186/1755-8794-2-66","title":"A metadata approach for clinical data management in translational genomics studies in breast cancer","year":2009,"lang":"en","type":"review","venue":"BMC Medical Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"NIHR Cambridge Biomedical Research Centre; Biomedical Research Council; National Institute for Health and Care Research; Engineering and Physical Sciences Research Council; Imperial Experimental Cancer Medicine Centre; Medical Research Council; Cancer Research UK","keywords":"Metadata; Human genetics; Genomics; Genome Biology; Breast cancer; Computational biology; Cancer; Biology; Personal genomics; Proteomics; Translational research; Bioinformatics; Data science; Computer science; Genetics; World Wide Web; Genome; Gene; Biotechnology","score_opus":0.4312372294638078,"score_gpt":0.5091070642886174,"score_spread":0.07786983482480958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116420111","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035508557,0.39833358,0.5555896,0.022911744,0.0019948806,0.0007913236,0.0019160948,0.001661033,0.013250829],"genre_scores_gemma":[0.036027394,0.2554413,0.69218534,0.0056759864,0.0012952599,0.0013225732,0.0046401373,0.00019715795,0.0032148617],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99300414,0.0034997638,0.0012385951,0.00058927963,0.0015345231,0.0001338179],"domain_scores_gemma":[0.9798076,0.013117877,0.0016982724,0.0021962663,0.0027992558,0.0003806299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020266267,0.0006920172,0.0013185957,0.010555358,0.00086462754,0.0044652103,0.003765004,0.0017758886,0.0013156778],"category_scores_gemma":[0.022117382,0.0005588261,0.001509203,0.017174391,0.0028916055,0.0082114795,0.0031051077,0.0032337862,0.0011374099],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007656562,0.00007569713,0.0025901198,0.007971228,0.0003047884,0.00051408407,0.001218368,0.0019990543,0.0020296357,0.10332422,0.02128565,0.8586106],"study_design_scores_gemma":[0.00006915546,0.000099078854,0.004177036,0.012485655,0.0005306289,0.0033697789,0.0014807115,0.00511589,0.004849735,0.14737156,0.8202898,0.00016109117],"about_ca_topic_score_codex":0.0035323042,"about_ca_topic_score_gemma":0.0035985147,"teacher_disagreement_score":0.020266267,"about_ca_system_score_codex":0.003770977,"about_ca_system_score_gemma":0.005778279,"threshold_uncertainty_score":0.10717952},"labels":[],"label_agreement":null},{"id":"W2116510544","doi":"10.1109/hicss.2011.21","title":"A Patient Profile Ontology in the Heterogeneous Domain of Complex and Chronic Health Conditions","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Ontology; Computer science; Vocabulary; Domain (mathematical analysis); Semantics (computer science); Semantic interoperability; Interoperability; Open Biomedical Ontologies; Controlled vocabulary; Domain knowledge; Terminology; Upper ontology; Knowledge management; Data science; Information retrieval; Ontology alignment; World Wide Web; Linguistics; Programming language","score_opus":0.03711540991911914,"score_gpt":0.2926917683378225,"score_spread":0.25557635841870335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2116510544","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02066147,0.0004854772,0.94156355,0.004861094,0.00025997098,0.001036576,0.0070799766,0.0014592053,0.022592742],"genre_scores_gemma":[0.14388506,0.0010280136,0.83770806,0.0010554085,0.00009590511,0.0007168754,0.009564088,0.00018048027,0.005766131],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99806863,0.0006484591,0.00040147235,0.0002880481,0.0004997696,0.00009378132],"domain_scores_gemma":[0.9975152,0.001278436,0.00026429867,0.00035381608,0.00043370467,0.00015450167],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004079029,0.00031096398,0.00039463036,0.0030199701,0.0012937592,0.002909048,0.00085816433,0.001177076,0.0023011796],"category_scores_gemma":[0.0055854106,0.00029931284,0.0011861919,0.0030733324,0.0010540964,0.0048803175,0.0018865884,0.0015004127,0.00049953006],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000873694,0.00017167826,0.0054598358,0.0005279622,0.00008416894,0.0017550823,0.0042470256,0.008680788,0.006232023,0.83595544,0.022857415,0.11394116],"study_design_scores_gemma":[0.000090339905,0.000080430706,0.006069497,0.0007791579,0.00020548004,0.0034167469,0.0027929326,0.08226564,0.005939016,0.26247305,0.63577956,0.00010815343],"about_ca_topic_score_codex":0.014787289,"about_ca_topic_score_gemma":0.017485285,"teacher_disagreement_score":0.014787289,"about_ca_system_score_codex":0.002196997,"about_ca_system_score_gemma":0.0062884325,"threshold_uncertainty_score":0.029402435},"labels":[],"label_agreement":null},{"id":"W2117770626","doi":"10.1186/2041-1480-2-s5-s11","title":"Assessment of NER solutions against the first and second CALBC Silver Standard Corpus","year":2011,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of New Brunswick","funders":"Concordia University; National Institute of Informatics; Universitat Jaume I; Magyar Tudományos Akadémia; Universidad Complutense de Madrid; Instituto de Salud Carlos III; Schweizerischer Nationalfonds zur Förderung der Wissenschaftlichen Forschung; Universiteit Antwerpen; Institute of Information Science, Academia Sinica; Academia Sinica; Fondazione Bruno Kessler; Universitat Pompeu Fabra","keywords":"Annotation; Computer science; Natural language processing; Set (abstract data type); Information retrieval; Artificial intelligence; Named-entity recognition; Identification (biology); Task (project management)","score_opus":0.024508395864490596,"score_gpt":0.2709703653760666,"score_spread":0.246461969511576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117770626","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4988708,0.009294843,0.11233407,0.009530649,0.005105942,0.0063997135,0.15175557,0.08018822,0.12652023],"genre_scores_gemma":[0.24761818,0.0011469098,0.18956491,0.0018583676,0.0004248533,0.0040957397,0.5196303,0.0067495047,0.028911136],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9848482,0.0052172565,0.0014591209,0.003269969,0.004435459,0.0007699557],"domain_scores_gemma":[0.9611128,0.015243142,0.001109315,0.005765111,0.015037908,0.0017317054],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016222341,0.0031493397,0.0017762132,0.007254868,0.0041315975,0.004464404,0.0043328307,0.004647211,0.012218383],"category_scores_gemma":[0.04008444,0.0008336405,0.0015025446,0.0046670446,0.0019046114,0.005292373,0.007011062,0.0035009591,0.012370092],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004235217,0.001981872,0.008345731,0.008969588,0.00071170286,0.0017448291,0.0056355232,0.019312264,0.05536585,0.008549973,0.5108233,0.37432414],"study_design_scores_gemma":[0.001932709,0.0022467785,0.07492264,0.0022556118,0.0006602399,0.0042278105,0.011551293,0.1829839,0.124296986,0.007877444,0.586273,0.0007716134],"about_ca_topic_score_codex":0.028848695,"about_ca_topic_score_gemma":0.035267398,"teacher_disagreement_score":0.028848695,"about_ca_system_score_codex":0.004647048,"about_ca_system_score_gemma":0.003354778,"threshold_uncertainty_score":0.08579296},"labels":[],"label_agreement":null},{"id":"W2117856677","doi":"10.2196/medinform.3387","title":"Enabling Online Studies of Conceptual Relationships Between Medical Terms: Developing an Efficient Web Platform","year":2014,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Unified Medical Language System; Computer science; Transitive relation; Set (abstract data type); Information retrieval; Data science; Programming language","score_opus":0.09734592737392041,"score_gpt":0.36913620903067895,"score_spread":0.27179028165675856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2117856677","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024018357,0.000268811,0.93658125,0.0008472347,0.00007798733,0.0011624967,0.0015045325,0.030104637,0.0054346803],"genre_scores_gemma":[0.07010151,0.00036702972,0.9208456,0.0002726463,0.000059217036,0.00080989784,0.0039107706,0.001535344,0.0020979084],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99663,0.001097195,0.00040394123,0.00056820875,0.0011005641,0.00020013023],"domain_scores_gemma":[0.98719573,0.0065666335,0.0009054347,0.0030695584,0.0014495524,0.00081315736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005863201,0.000938787,0.0008499338,0.0032303727,0.0007695532,0.003925226,0.002683699,0.0011929338,0.0060681454],"category_scores_gemma":[0.013048935,0.0007384022,0.0016585768,0.0021226641,0.0009379744,0.007816033,0.0055724694,0.0014380977,0.0033553871],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001432334,0.0014938739,0.012835251,0.0018404453,0.0003155814,0.003037907,0.002888999,0.02166776,0.08185341,0.099263854,0.03750996,0.73586065],"study_design_scores_gemma":[0.0007227077,0.00064774696,0.0056144344,0.00052390876,0.0002745095,0.0018314827,0.0019771187,0.40484446,0.10276435,0.16575831,0.31470045,0.00034049124],"about_ca_topic_score_codex":0.0020259493,"about_ca_topic_score_gemma":0.0015677074,"teacher_disagreement_score":0.0060681454,"about_ca_system_score_codex":0.00094854983,"about_ca_system_score_gemma":0.0028605505,"threshold_uncertainty_score":0.031007946},"labels":[],"label_agreement":null},{"id":"W2118284237","doi":"10.1093/database/bas036","title":"Recent advances in biocuration: Meeting Report from the fifth International Biocuration Conference","year":2012,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Human Genome Research Institute; George Washington University; Wellcome Trust; Georgetown University; Institute of Genetics; University of Oxford; Ontario Institute for Cancer Research","keywords":"Library science; World Wide Web; Promotion (chess); Political science; Computer science","score_opus":0.03594583616335801,"score_gpt":0.3235142753631416,"score_spread":0.2875684391997836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2118284237","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010091448,0.6386249,0.006999975,0.23352784,0.047366727,0.00019078929,0.0043178704,0.0008196793,0.058060717],"genre_scores_gemma":[0.087724276,0.67366594,0.03171231,0.053782646,0.031095084,0.0005180248,0.02551178,0.001403174,0.09458679],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9939355,0.0013078229,0.0005917962,0.00072471227,0.0024537705,0.0009863769],"domain_scores_gemma":[0.98013353,0.0031167723,0.0014112083,0.00086629,0.008069858,0.0064023724],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.02507118,0.0012143962,0.0008778839,0.004743088,0.001611072,0.009187875,0.0021180161,0.0039744214,0.017397212],"category_scores_gemma":[0.014121438,0.00046540002,0.0014143209,0.008806737,0.0012780653,0.006810134,0.0059801457,0.0032529968,0.008849174],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019743554,0.000115142284,0.0031529283,0.0018752407,0.00008266776,0.00018593029,0.0007155735,0.00031476215,0.0012170626,0.004606392,0.755484,0.23205289],"study_design_scores_gemma":[0.000007956158,0.000040964183,0.003126605,0.00061674987,0.000042625354,0.00011211697,0.00056391925,0.00008205102,0.00074920245,0.0006378872,0.99399614,0.000023706152],"about_ca_topic_score_codex":0.006504679,"about_ca_topic_score_gemma":0.0088802185,"teacher_disagreement_score":0.9908121,"about_ca_system_score_codex":0.003259841,"about_ca_system_score_gemma":0.007861492,"threshold_uncertainty_score":0.13259065},"labels":[],"label_agreement":null},{"id":"W2119707622","doi":"10.1109/hicss.2002.994069","title":"Using the XML-based Clinical Document Architecture for exchange of structured discharge summaries","year":2003,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"XML; Computer science; Architecture; Health care; Markup language; World Wide Web; Medicine","score_opus":0.05588777479120789,"score_gpt":0.3647795885601137,"score_spread":0.30889181376890584,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2119707622","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035773886,0.0003992268,0.9745668,0.0021793093,0.00021025017,0.000840704,0.0009797756,0.01015579,0.007090911],"genre_scores_gemma":[0.037265673,0.00067220913,0.9485535,0.0007113675,0.000113107664,0.0008155184,0.0038603581,0.0009787808,0.007029471],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9887235,0.0050131227,0.0027668674,0.0009437393,0.0022498444,0.00030300757],"domain_scores_gemma":[0.97325844,0.0120665645,0.0019765587,0.0060087154,0.0057226853,0.0009670358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019356959,0.0007437018,0.00061173894,0.0037553534,0.0018727838,0.009725494,0.0025312342,0.0032554793,0.0058937417],"category_scores_gemma":[0.03494729,0.0012272181,0.0011745772,0.0044623367,0.001827709,0.007512877,0.0042191185,0.0030756416,0.0047012586],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065425393,0.00040832948,0.005511549,0.0016241473,0.00024964748,0.0017163112,0.0071014804,0.029562498,0.020276949,0.39316922,0.05691187,0.4828137],"study_design_scores_gemma":[0.00048746954,0.00049837824,0.0028084582,0.0012838355,0.00022736657,0.0025575617,0.0008231896,0.121529736,0.025152922,0.15294987,0.69129956,0.00038163728],"about_ca_topic_score_codex":0.007634894,"about_ca_topic_score_gemma":0.0055341413,"teacher_disagreement_score":0.019356959,"about_ca_system_score_codex":0.0026172304,"about_ca_system_score_gemma":0.005387697,"threshold_uncertainty_score":0.10237056},"labels":[],"label_agreement":null},{"id":"W2120235270","doi":"10.1002/j.0022-0337.2011.75.1.tb05024.x","title":"The Development of a Dental Diagnostic Terminology","year":2011,"lang":"en","type":"article","venue":"Journal of Dental Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Institute of Dental and Craniofacial Research","keywords":"Terminology; Medical diagnosis; Standardization; Quality assurance; Work (physics); Medicine; Medical physics; Medical education; Computer science; Pathology; Linguistics","score_opus":0.021589382400888018,"score_gpt":0.294003046089091,"score_spread":0.27241366368820297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120235270","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03224279,0.0042304154,0.8789291,0.019972047,0.0030643553,0.007804718,0.0060438267,0.0013274336,0.04638528],"genre_scores_gemma":[0.036445722,0.0015621199,0.9511862,0.001146687,0.00021750698,0.0024985278,0.004670708,0.0002293713,0.0020432451],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9704195,0.009419945,0.009347472,0.0019897341,0.00807828,0.00074518076],"domain_scores_gemma":[0.9509561,0.014690467,0.0040436536,0.0042977696,0.024606796,0.0014053192],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030675545,0.00082687946,0.0010939769,0.016764458,0.003788566,0.0067301895,0.0038382972,0.0019430565,0.0031635554],"category_scores_gemma":[0.062546656,0.00081058114,0.0021676258,0.011930655,0.0034358762,0.0077681844,0.0063622985,0.004210937,0.0020559367],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013916893,0.00021106831,0.008474707,0.0033747558,0.00010057169,0.0009429797,0.01569993,0.0025126815,0.012115983,0.3510736,0.05386219,0.55149233],"study_design_scores_gemma":[0.00008425808,0.0002195838,0.009007234,0.005156152,0.00023696749,0.0026100867,0.013519849,0.01076914,0.01025303,0.14128785,0.8066068,0.0002489483],"about_ca_topic_score_codex":0.0073866234,"about_ca_topic_score_gemma":0.005122983,"teacher_disagreement_score":0.030675545,"about_ca_system_score_codex":0.00815998,"about_ca_system_score_gemma":0.029668491,"threshold_uncertainty_score":0.16222972},"labels":[],"label_agreement":null},{"id":"W2120271288","doi":"10.1016/j.ympev.2008.08.027","title":"Correctly assigning original discoveries to original authors","year":2008,"lang":"en","type":"letter","venue":"Molecular Phylogenetics and Evolution","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Allison University","funders":"","keywords":"Biology; Evolutionary biology; Computational biology","score_opus":0.012279437898530542,"score_gpt":0.25059385374857296,"score_spread":0.23831441585004243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2120271288","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029087767,0.0014126971,0.009714633,0.876359,0.10306676,0.000045578465,0.00032254678,0.00038061265,0.005789465],"genre_scores_gemma":[0.072989464,0.0061154896,0.035189133,0.6040912,0.24694766,0.00028571038,0.0007174844,0.0006215139,0.03304225],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9904482,0.0035490708,0.0017426773,0.0011316774,0.0023867623,0.0007415824],"domain_scores_gemma":[0.8909861,0.060324837,0.0056658015,0.012758914,0.026266746,0.0039975974],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.014447873,0.0007300945,0.001524165,0.0022748369,0.002560425,0.004963392,0.0019692446,0.015169967,0.0063070115],"category_scores_gemma":[0.16518226,0.0008775766,0.0011696395,0.0021994987,0.0028987993,0.0047978237,0.0022197012,0.01795008,0.008129545],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031679246,0.000042711632,0.0021442194,0.00023410328,0.00007728481,0.0043046093,0.0002693017,0.00036549105,0.0009706194,0.009202177,0.89605385,0.08601884],"study_design_scores_gemma":[0.0002923954,0.00008664302,0.0026815862,0.00047340032,0.00030319818,0.015825052,0.00044320332,0.0064799185,0.0055727577,0.108152576,0.8595353,0.00015395758],"about_ca_topic_score_codex":0.0009648968,"about_ca_topic_score_gemma":0.0015602554,"teacher_disagreement_score":0.98555213,"about_ca_system_score_codex":0.0027217332,"about_ca_system_score_gemma":0.003799737,"threshold_uncertainty_score":0.076408565},"labels":[],"label_agreement":null},{"id":"W2121237412","doi":"10.1093/nar/gkq1173","title":"Towards BioDBcore: a community-defined information specification for biological databases","year":2010,"lang":"en","type":"editorial","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Institute of General Medical Sciences; Natural Environment Research Council; National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council; National Institutes of Health","keywords":"Interoperability; Scope (computer science); Relevance (law); Consistency (knowledge bases); Database; Resource (disambiguation); Biology; Semantic heterogeneity; Knowledge management; Data science; Computer science; World Wide Web; Semantic Web","score_opus":0.12545575823948255,"score_gpt":0.40574465489601524,"score_spread":0.28028889665653267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121237412","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00039007954,0.017270403,0.04469205,0.14737064,0.78413683,0.00020455592,0.00029787334,0.0005216589,0.005115935],"genre_scores_gemma":[0.0052740145,0.03616327,0.06183384,0.16589573,0.68791276,0.0007539253,0.001192701,0.0013050348,0.039668668],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9876167,0.0032521086,0.0022088736,0.0009110775,0.0056013553,0.00040982032],"domain_scores_gemma":[0.95878726,0.018425671,0.0021691443,0.001978956,0.015287847,0.003351141],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.028583277,0.0016472476,0.001584624,0.0032681185,0.0021155786,0.009109698,0.0043080803,0.013596858,0.0027750547],"category_scores_gemma":[0.034977827,0.0012971739,0.002124682,0.0021906078,0.0054014795,0.011903704,0.004349228,0.034786157,0.003991825],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000061467545,0.000036101323,0.00007681495,0.00071330095,0.00004131556,0.00037136817,0.00041314046,0.00037727234,0.0009909283,0.032116175,0.93147635,0.03332558],"study_design_scores_gemma":[0.000023383369,0.000011573995,0.000055670003,0.00029837294,0.000014635437,0.00015803588,0.00004139443,0.00026670913,0.00022402815,0.0043326374,0.9945533,0.000020181924],"about_ca_topic_score_codex":0.0029511745,"about_ca_topic_score_gemma":0.0056182034,"teacher_disagreement_score":0.9908903,"about_ca_system_score_codex":0.0046818443,"about_ca_system_score_gemma":0.0070610796,"threshold_uncertainty_score":0.15116459},"labels":[],"label_agreement":null},{"id":"W2121981832","doi":"10.1186/1471-2105-13-s9-s2","title":"Modeling and mining term association for improving biomedical information retrieval performance","year":2012,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Information retrieval; Ranking (information retrieval); Term (time); Term Discrimination; Set (abstract data type); Index (typography); Sentence; Association rule learning; Data mining; Paragraph; Baseline (sea); Artificial intelligence; Search engine; Concept search; Web search query; World Wide Web","score_opus":0.01922306671104155,"score_gpt":0.25661114349442515,"score_spread":0.2373880767833836,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2121981832","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2823738,0.006230456,0.70120287,0.0006889895,0.00018141096,0.00034539137,0.0014954676,0.005190853,0.0022907464],"genre_scores_gemma":[0.72882843,0.0012786975,0.2641702,0.00015501284,0.00024902652,0.00025427932,0.0035089394,0.00020214575,0.001353249],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99760336,0.0007065035,0.00029978354,0.00066710694,0.00052730594,0.00019599669],"domain_scores_gemma":[0.99451435,0.0028951012,0.0007948313,0.0005432714,0.0011046133,0.00014789445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029491463,0.0014076835,0.0016889755,0.0047716885,0.00067341287,0.0012455389,0.0012760704,0.0010275949,0.0008680302],"category_scores_gemma":[0.010563579,0.0003167765,0.0017758819,0.005065153,0.0004661964,0.0023683605,0.0006922557,0.0011017265,0.0011827636],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013362985,0.00104634,0.04938056,0.00078767847,0.00055360515,0.0003799863,0.00034916785,0.21372004,0.045995012,0.003292799,0.008934731,0.6742238],"study_design_scores_gemma":[0.000032470354,0.0002765025,0.005062254,0.000020861704,0.0001886033,0.00022242809,0.00005476627,0.9807407,0.009129659,0.0025100387,0.0017243141,0.000037399357],"about_ca_topic_score_codex":0.007986033,"about_ca_topic_score_gemma":0.008732699,"teacher_disagreement_score":0.007986033,"about_ca_system_score_codex":0.0009777248,"about_ca_system_score_gemma":0.001911672,"threshold_uncertainty_score":0.015879095},"labels":[],"label_agreement":null},{"id":"W2122062351","doi":"10.1093/bioinformatics/btr389","title":"BRISK—research-oriented storage kit for biology-related data","year":2011,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Java; Documentation; Data sharing; Data management; Software; Data access; Web application; World Wide Web; Data science; Database; Operating system","score_opus":0.17065333150659484,"score_gpt":0.373670698899958,"score_spread":0.20301736739336318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122062351","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010532199,0.0015856702,0.27345476,0.0025804546,0.0008406073,0.0018654127,0.15422761,0.5300828,0.024830492],"genre_scores_gemma":[0.07235617,0.0018726273,0.47585046,0.0026313437,0.0006158757,0.0044354303,0.3565519,0.05262035,0.033065826],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9962586,0.0006432381,0.00079000276,0.0006444556,0.00136551,0.00029812925],"domain_scores_gemma":[0.9794049,0.0047652507,0.0025381814,0.0070306323,0.0045251446,0.0017357699],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00783029,0.001786459,0.0016122236,0.0052083475,0.0010865725,0.0057658064,0.0051324526,0.0010475068,0.04806086],"category_scores_gemma":[0.01787144,0.0014717164,0.0011162652,0.006812205,0.0010034245,0.007183745,0.005653737,0.0020607589,0.07470085],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027811965,0.00023536543,0.0052507916,0.00368831,0.00027815625,0.0006383753,0.0013690756,0.0019013591,0.026054226,0.016422411,0.66867745,0.27270323],"study_design_scores_gemma":[0.000796615,0.00040815555,0.00835936,0.0008435685,0.0002519446,0.0010360224,0.0005440875,0.015682325,0.059029933,0.025012923,0.8876004,0.00043466248],"about_ca_topic_score_codex":0.0020084116,"about_ca_topic_score_gemma":0.0014229165,"teacher_disagreement_score":0.04806086,"about_ca_system_score_codex":0.0013588772,"about_ca_system_score_gemma":0.0031262033,"threshold_uncertainty_score":0.1607796},"labels":[],"label_agreement":null},{"id":"W2122174694","doi":"","title":"The National Center for Biomedical Ontology: Advancing Biomedicine through Structured \\nOrganization of Scientific Knowledge","year":2006,"lang":"en","type":"review","venue":"eScholarship (California Digital Library)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":167,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ontology; Dissemination; Biomedicine; Context (archaeology); Computer science; Knowledge management; Data science; Resource (disambiguation); Open Biomedical Ontologies; World Wide Web; Upper ontology; Bioinformatics; Domain knowledge; Suggested Upper Merged Ontology","score_opus":0.028211212614483148,"score_gpt":0.30672890902494043,"score_spread":0.27851769641045726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2122174694","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026913479,0.005836787,0.8477725,0.05910194,0.00399612,0.0013151176,0.0051899105,0.01006384,0.064032406],"genre_scores_gemma":[0.014312164,0.0062032067,0.94460744,0.0086779725,0.00091969233,0.0011659967,0.013575474,0.0015136506,0.009024429],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98107374,0.0077907494,0.0027705035,0.0022237094,0.0052617304,0.000879588],"domain_scores_gemma":[0.96300006,0.012489212,0.0027957226,0.01216428,0.0061640753,0.0033866204],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.034630187,0.0013781289,0.0015156067,0.009176549,0.004924847,0.014988557,0.005264577,0.0043056817,0.008441823],"category_scores_gemma":[0.04336624,0.0013771428,0.0027003644,0.009730598,0.007878216,0.020508492,0.021776067,0.0082319155,0.004824776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009271112,0.0001405949,0.0013835534,0.0012474764,0.00010893309,0.00019707033,0.0023208726,0.0013504185,0.0029917613,0.6447626,0.1681808,0.17722327],"study_design_scores_gemma":[0.000065216416,0.000036557733,0.0007687139,0.0011638174,0.00006367895,0.00015790871,0.00062679686,0.0059606037,0.0017351409,0.27723664,0.7121144,0.000070536466],"about_ca_topic_score_codex":0.022108594,"about_ca_topic_score_gemma":0.0202988,"teacher_disagreement_score":0.98501146,"about_ca_system_score_codex":0.0073014675,"about_ca_system_score_gemma":0.036317214,"threshold_uncertainty_score":0.18314415},"labels":[],"label_agreement":null},{"id":"W2123075577","doi":"10.1007/978-3-642-13059-5_33","title":"Using Classifier Performance Visualization to Improve Collective Ranking Techniques for Biomedical Abstracts Classification","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; University of New Brunswick","funders":"","keywords":"Computer science; Classifier (UML); Cluster analysis; Workload; Machine learning; Artificial intelligence; Data mining; Visualization; Ranking (information retrieval)","score_opus":0.03871160874918192,"score_gpt":0.32331913398872597,"score_spread":0.284607525239544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123075577","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08883728,0.0025009206,0.8344015,0.002095987,0.0010846091,0.00029909628,0.0033318824,0.06051079,0.0069379346],"genre_scores_gemma":[0.40091023,0.0006277398,0.5847431,0.00021328614,0.00056254055,0.00029318332,0.004248482,0.0029220271,0.005479389],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966478,0.0008808819,0.0003708936,0.00041794888,0.0014639652,0.00021849232],"domain_scores_gemma":[0.9749552,0.011536337,0.0017290335,0.0025972938,0.008506925,0.0006752875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063660857,0.0019264945,0.001760929,0.007354068,0.0011033902,0.00473658,0.0014694849,0.0013697566,0.006163057],"category_scores_gemma":[0.025612619,0.0006096148,0.0011295845,0.0056364085,0.000320506,0.0038450947,0.0017882822,0.0024052272,0.0030485434],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054149836,0.00029515507,0.008754197,0.00032514683,0.00023829803,0.00013026063,0.0005671846,0.017146785,0.01535942,0.003935739,0.04806835,0.90463793],"study_design_scores_gemma":[0.00014443665,0.0003631892,0.007964532,0.00014351397,0.0002905156,0.00029023646,0.0002825517,0.9176525,0.03229102,0.022624807,0.017793953,0.00015882518],"about_ca_topic_score_codex":0.0035657976,"about_ca_topic_score_gemma":0.0058827694,"teacher_disagreement_score":0.007354068,"about_ca_system_score_codex":0.00084718387,"about_ca_system_score_gemma":0.0011301839,"threshold_uncertainty_score":0.033667505},"labels":[],"label_agreement":null},{"id":"W2123574488","doi":"10.1186/2041-1480-3-6","title":"Extending and encoding existing biological terminologies and datasets for use in the reasoned semantic web","year":2012,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; St. Paul's Hospital","funders":"Canadian Institutes of Health Research; Heart and Stroke Foundation of British Columbia and Yukon; Microsoft Research; Natural Sciences and Engineering Research Council of Canada; Heart and Stroke Foundation of Canada","keywords":"Computer science; Data science; Context (archaeology); Information retrieval; Artificial intelligence; Data mining","score_opus":0.0951616708167776,"score_gpt":0.3419260940137993,"score_spread":0.24676442319702172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2123574488","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010774123,0.00038291115,0.9668327,0.0029115917,0.0002596452,0.0004589593,0.00810029,0.0059301206,0.0043497407],"genre_scores_gemma":[0.071267694,0.00076079543,0.90296876,0.0010157932,0.000086165564,0.0005822805,0.021362925,0.0009497544,0.0010058156],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913147,0.0024571952,0.0020502224,0.00128571,0.0025938987,0.0002983433],"domain_scores_gemma":[0.970713,0.011995275,0.0021526904,0.010554578,0.004002335,0.0005820468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014877215,0.0009751493,0.0008893709,0.005668191,0.0017848875,0.008276316,0.003171911,0.0020161604,0.0017131544],"category_scores_gemma":[0.032808952,0.0007252061,0.0030851245,0.006501879,0.0029996976,0.012397952,0.0065688626,0.0049411985,0.0011111393],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045279064,0.00048997765,0.013587015,0.0023246608,0.00041270073,0.002037396,0.006766543,0.051581327,0.010602086,0.6156641,0.02731694,0.26876447],"study_design_scores_gemma":[0.00008603672,0.00008192976,0.002527317,0.0011113882,0.000260067,0.0006084524,0.0016940943,0.13401194,0.014493628,0.609397,0.23553818,0.00018989094],"about_ca_topic_score_codex":0.008531011,"about_ca_topic_score_gemma":0.011180525,"teacher_disagreement_score":0.014877215,"about_ca_system_score_codex":0.0029605832,"about_ca_system_score_gemma":0.0070372294,"threshold_uncertainty_score":0.078679144},"labels":[],"label_agreement":null},{"id":"W2124057135","doi":"10.1089/bio.2011.0020","title":"A Proposed Schema for Classifying Human Research Biobanks","year":2011,"lang":"en","type":"article","venue":"Biopreservation and Biobanking","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"BC Cancer Agency; Michael Smith Health Research BC","keywords":"Biobank; Schema (genetic algorithms); Computer science; Data science; Computational biology; Information retrieval; Bioinformatics; Biology","score_opus":0.4782123450186061,"score_gpt":0.4320864660738816,"score_spread":0.046125878944724474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124057135","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009208681,0.0006661725,0.93721455,0.006822676,0.0007296309,0.004456509,0.011092768,0.004511163,0.025297865],"genre_scores_gemma":[0.016707703,0.00050679664,0.96229774,0.0009924129,0.00013272704,0.0015524017,0.0141836135,0.00022213801,0.0034045768],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97550493,0.0074790698,0.007902102,0.0038437536,0.004125071,0.0011450644],"domain_scores_gemma":[0.96278733,0.010896717,0.0035398952,0.0075617107,0.012955026,0.0022592705],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02764254,0.0012181979,0.0012052658,0.011349803,0.004212356,0.018344482,0.005391628,0.0046797097,0.00721],"category_scores_gemma":[0.0288257,0.0013650234,0.0030793822,0.013415128,0.0036175377,0.01834007,0.004900276,0.0036978368,0.004598032],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000132386,0.00027228738,0.0103996275,0.00095839205,0.00011394736,0.0008157733,0.006731838,0.0050070877,0.0029607508,0.7880224,0.051025607,0.13355985],"study_design_scores_gemma":[0.000077542296,0.0001289677,0.0034592731,0.0014538523,0.0001542376,0.0016145465,0.005547512,0.026012558,0.003657116,0.18893333,0.7687902,0.00017094197],"about_ca_topic_score_codex":0.014092782,"about_ca_topic_score_gemma":0.009515945,"teacher_disagreement_score":0.02764254,"about_ca_system_score_codex":0.0052901236,"about_ca_system_score_gemma":0.013117314,"threshold_uncertainty_score":0.14618945},"labels":[],"label_agreement":null},{"id":"W2124308436","doi":"10.1186/2041-1480-4-s1-s1","title":"Ontology-Based Querying with Bio2RDF’s Linked Open Data","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Scripting language; Ontology; Data integration; Linked data; Information retrieval; World Wide Web; Vocabulary; Semantic Web; Data science; Database; Data mining","score_opus":0.04962742204416401,"score_gpt":0.32060276404958415,"score_spread":0.27097534200542017,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124308436","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011733137,0.00052170805,0.7851255,0.0027135224,0.00044478255,0.0008304507,0.04106994,0.1433717,0.014189314],"genre_scores_gemma":[0.066318564,0.00082223176,0.7277395,0.001596675,0.00015178532,0.0012452588,0.1777013,0.019441882,0.0049828],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9909869,0.0022299616,0.0014557764,0.0015967059,0.0033196271,0.00041111896],"domain_scores_gemma":[0.98742604,0.0047804383,0.00068876165,0.004763207,0.0017774267,0.000564124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01759832,0.0014511655,0.0010986405,0.005941531,0.0020306169,0.007899354,0.0038787685,0.0021487826,0.006828028],"category_scores_gemma":[0.028841615,0.001033434,0.0040651616,0.005105175,0.0019910745,0.008614293,0.0074698045,0.0029694638,0.003756909],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015633011,0.00093097467,0.017794468,0.0028368724,0.00085153215,0.0028425795,0.005323883,0.048642825,0.030799575,0.3127915,0.30245492,0.27316755],"study_design_scores_gemma":[0.00026332933,0.00010454248,0.00465359,0.0007246576,0.00014440683,0.000863779,0.000873754,0.12744485,0.032756116,0.1639405,0.66786397,0.00036648745],"about_ca_topic_score_codex":0.018368859,"about_ca_topic_score_gemma":0.018356962,"teacher_disagreement_score":0.018368859,"about_ca_system_score_codex":0.0048141186,"about_ca_system_score_gemma":0.004494258,"threshold_uncertainty_score":0.09306991},"labels":[],"label_agreement":null},{"id":"W2124714582","doi":"10.1093/bioinformatics/btr452","title":"OrganismTagger: detection, normalization and grounding of organism entities in biomedical documents","year":2011,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of New Brunswick","funders":"","keywords":"Computer science; Organism; Information retrieval; Taxonomy (biology); Natural language processing; Precision and recall; Named-entity recognition; Artificial intelligence; Task (project management); Biology","score_opus":0.014569750501862155,"score_gpt":0.2312963566121826,"score_spread":0.21672660611032044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2124714582","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11555527,0.007690686,0.50466543,0.0018104544,0.00084349664,0.0011761705,0.08185356,0.2745901,0.0118147135],"genre_scores_gemma":[0.101216905,0.0015020831,0.80899864,0.0004985353,0.00014659272,0.00045683794,0.07945493,0.0025699562,0.005155517],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99719113,0.00038859848,0.00050452666,0.0010622572,0.0007568059,0.00009676625],"domain_scores_gemma":[0.99368095,0.0028428682,0.0011430964,0.0008884022,0.0012541712,0.00019044994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002935746,0.0013054307,0.00086773164,0.012545528,0.00083969673,0.0028812701,0.0016621384,0.0017453092,0.0043003885],"category_scores_gemma":[0.010837156,0.00073019817,0.00085301726,0.005605407,0.0007813099,0.0037207047,0.0019180411,0.0008375224,0.006540365],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048616872,0.00029278835,0.029108413,0.0046235407,0.00036060184,0.0014566985,0.0014503995,0.0041622547,0.09488849,0.0047795773,0.08869346,0.76969755],"study_design_scores_gemma":[0.0002519022,0.00064381206,0.08641172,0.0021945303,0.0007810085,0.007306713,0.0017553314,0.111499116,0.33558607,0.015700804,0.4373908,0.0004781778],"about_ca_topic_score_codex":0.0040796353,"about_ca_topic_score_gemma":0.0065145013,"teacher_disagreement_score":0.012545528,"about_ca_system_score_codex":0.0011150924,"about_ca_system_score_gemma":0.0022583953,"threshold_uncertainty_score":0.015525937},"labels":[],"label_agreement":null},{"id":"W2125160887","doi":"10.1186/1742-5581-3-11","title":"LitMiner: integration of library services within a bio-informatics application","year":2006,"lang":"en","type":"article","venue":"Biomedical Digital Libraries","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Digital library; World Wide Web; Computer science; Service (business); Suite; Process (computing); Subject (documents); Order (exchange); Information retrieval; Data science","score_opus":0.0053830339766246585,"score_gpt":0.20749004497913387,"score_spread":0.20210701100250922,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2125160887","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05262056,0.0015822484,0.5991958,0.0028136442,0.00013746289,0.000961062,0.004346435,0.31733143,0.02101145],"genre_scores_gemma":[0.19092996,0.0014987664,0.77325535,0.0019293423,0.0001264493,0.0007354725,0.011131776,0.008483348,0.011909583],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963174,0.0011965461,0.0005245536,0.0005636291,0.0012313339,0.00016646758],"domain_scores_gemma":[0.9919076,0.0046330434,0.00056908163,0.0014415152,0.0010080712,0.00044070373],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0053874035,0.00081174716,0.0007808015,0.0047324575,0.0010859827,0.0039498596,0.0022592386,0.0013720662,0.0062389467],"category_scores_gemma":[0.014279282,0.00060369185,0.001182962,0.004320637,0.00069470954,0.005589912,0.004187868,0.0009924519,0.003994776],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018631943,0.0007649546,0.010909323,0.0030024939,0.00043834472,0.0024623303,0.0047556153,0.0061291545,0.027984353,0.021982841,0.0709921,0.8487153],"study_design_scores_gemma":[0.00039854128,0.0005824891,0.011250946,0.00090981036,0.00045477963,0.0047767293,0.0017811677,0.1475342,0.07861212,0.028417373,0.7247989,0.00048282812],"about_ca_topic_score_codex":0.0036681916,"about_ca_topic_score_gemma":0.0042826594,"teacher_disagreement_score":0.9960501,"about_ca_system_score_codex":0.0015474646,"about_ca_system_score_gemma":0.0021154587,"threshold_uncertainty_score":0.028491676},"labels":[],"label_agreement":null},{"id":"W2126122334","doi":"10.1093/bioinformatics/bts542","title":"Application and evaluation of automated methods to extract neuroanatomical connectivity statements from free text","year":2012,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of General Medical Sciences; Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; National Institutes of Health; Michael Smith Health Research BC","keywords":"Computer science; Artificial intelligence; Natural language processing; Text messaging; Software; Machine learning; Data mining; Programming language; World Wide Web","score_opus":0.045433913833232746,"score_gpt":0.4067668540895707,"score_spread":0.36133294025633794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126122334","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6121725,0.008380027,0.23130447,0.0017228756,0.00078425417,0.003879971,0.051454697,0.07585865,0.014442537],"genre_scores_gemma":[0.545368,0.0014340085,0.3627898,0.00034198005,0.00026859937,0.0018439628,0.08114684,0.0013850927,0.0054217707],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98933804,0.0040067476,0.0016276981,0.0023343265,0.0022569145,0.00043642247],"domain_scores_gemma":[0.92059803,0.05798049,0.0031490629,0.0051418566,0.011964372,0.0011661239],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012044203,0.002793921,0.0010839867,0.012336367,0.0012099347,0.0030692199,0.0024078588,0.002136392,0.005069456],"category_scores_gemma":[0.04989831,0.0006449429,0.0012942725,0.0046275365,0.0007519549,0.004391008,0.0024821481,0.0011061161,0.0040469086],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028410077,0.002073376,0.0303071,0.006221988,0.0012681687,0.0012231554,0.0023415869,0.02348241,0.03236118,0.002326249,0.05888581,0.8366679],"study_design_scores_gemma":[0.0014187525,0.0025622647,0.10373283,0.0010257419,0.0010923767,0.002906439,0.002846502,0.72529274,0.10060745,0.006951152,0.051148985,0.00041480132],"about_ca_topic_score_codex":0.011765667,"about_ca_topic_score_gemma":0.010795158,"teacher_disagreement_score":0.012336367,"about_ca_system_score_codex":0.0017169318,"about_ca_system_score_gemma":0.00315281,"threshold_uncertainty_score":0.06369662},"labels":[],"label_agreement":null},{"id":"W2126320151","doi":"10.1093/bioinformatics/btl410","title":"IndexToolkit: an open source toolbox to index protein databases for high-throughput proteomics","year":2006,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Search engine indexing; Database; Toolbox; MIT License; Interface (matter); Application programming interface; Information retrieval; Source code; Software; Open source; Database index; License; Database search engine; Index (typography); World Wide Web; Search engine; Programming language; Operating system","score_opus":0.03160131351498109,"score_gpt":0.3029832167974074,"score_spread":0.27138190328242634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2126320151","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005852788,0.0032671513,0.3961022,0.0004530502,0.0004475997,0.00084704463,0.114334084,0.46922776,0.009468315],"genre_scores_gemma":[0.022698727,0.003420784,0.4775941,0.00066523725,0.00021743368,0.0025922605,0.4166763,0.06384082,0.01229431],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99883956,0.00013389217,0.00018083815,0.00022329297,0.0005003689,0.00012202024],"domain_scores_gemma":[0.9976501,0.0008386877,0.00033582497,0.00044287994,0.00048123221,0.0002512741],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021038575,0.0035081794,0.0024019,0.007614188,0.0016228016,0.004559374,0.0049509285,0.0013340616,0.0353654],"category_scores_gemma":[0.007732161,0.0020236415,0.0017538471,0.008522292,0.0006652536,0.0056152577,0.0040719337,0.0022710778,0.04268356],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011526111,0.00021200845,0.0017909347,0.0057348646,0.0003548979,0.00082730665,0.0005089535,0.00240176,0.034075186,0.009018668,0.63958806,0.3043348],"study_design_scores_gemma":[0.00068860856,0.00025277855,0.005993571,0.00091395614,0.00024085521,0.0026375384,0.00024257695,0.0352215,0.07611455,0.032058354,0.8451453,0.0004904471],"about_ca_topic_score_codex":0.0022968757,"about_ca_topic_score_gemma":0.0023145988,"teacher_disagreement_score":0.0353654,"about_ca_system_score_codex":0.0012759792,"about_ca_system_score_gemma":0.0022113642,"threshold_uncertainty_score":0.11830908},"labels":[],"label_agreement":null},{"id":"W2128119041","doi":"10.1186/1758-2946-4-9","title":"A chemical specialty semantic network for the Unified Medical Language System","year":2012,"lang":"en","type":"article","venue":"Journal of Cheminformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"New York Institute of Technology","funders":"","keywords":"Unified Medical Language System; Computer science; Ontology; Semantics (computer science); Semantic network; Natural language processing; Information retrieval; Artificial intelligence; Data mining; Programming language","score_opus":0.012390122717380214,"score_gpt":0.2734242568030976,"score_spread":0.2610341340857174,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128119041","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036695845,0.00046838325,0.92996615,0.0016274519,0.00008808133,0.00078158744,0.0061626886,0.0029531366,0.021256618],"genre_scores_gemma":[0.1411197,0.00031357925,0.8491815,0.00021951976,0.00003888755,0.0006213902,0.0068354285,0.0001600506,0.0015098507],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99756145,0.00087672117,0.00023163401,0.00043292268,0.00080447196,0.000092790615],"domain_scores_gemma":[0.9959741,0.002451974,0.00032678002,0.00034701728,0.0007852844,0.000114886636],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003177079,0.00039765125,0.00037574957,0.0045696977,0.0012657655,0.0021652777,0.000878936,0.0010391185,0.006542001],"category_scores_gemma":[0.013613797,0.00027134572,0.00097539544,0.0044894903,0.0010088529,0.0033806746,0.0016616953,0.0007101985,0.0011330697],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035413142,0.00013534905,0.00976141,0.0012080373,0.000119920165,0.00062703184,0.0019454266,0.119667366,0.006307434,0.5733807,0.021405853,0.2650873],"study_design_scores_gemma":[0.000064180895,0.00007517538,0.0047515435,0.00043034583,0.00016255608,0.00047879765,0.00068901543,0.5312783,0.005527392,0.33878276,0.117696814,0.00006316044],"about_ca_topic_score_codex":0.0090535525,"about_ca_topic_score_gemma":0.0137745105,"teacher_disagreement_score":0.0090535525,"about_ca_system_score_codex":0.003766058,"about_ca_system_score_gemma":0.003957933,"threshold_uncertainty_score":0.027324796},"labels":[],"label_agreement":null},{"id":"W2128184656","doi":"10.1007/s10278-013-9599-2","title":"Quantitative Imaging Biomarker Ontology (QIBO) for Knowledge Representation of Biomedical Imaging Biomarkers","year":2013,"lang":"en","type":"review","venue":"Journal of Digital Imaging","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Research (Canada)","funders":"National Cancer Institute; National Institute of Standards and Technology; U.S. Department of Commerce","keywords":"Computer science; Biomarker; Imaging biomarker; Biomarker discovery; Ontology; Medical imaging; Artificial intelligence; Medical physics; Data science; Magnetic resonance imaging; Medicine; Radiology; Proteomics","score_opus":0.08355144632235116,"score_gpt":0.4109724050661386,"score_spread":0.3274209587437874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128184656","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005084865,0.58289146,0.38185126,0.0086671095,0.0011593389,0.0006602911,0.0061600017,0.002127147,0.011398515],"genre_scores_gemma":[0.030302295,0.53959435,0.40425408,0.003936068,0.0004680841,0.0007442666,0.015003958,0.00025399646,0.0054429565],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988663,0.00018946419,0.00019522524,0.00019520774,0.0004868967,0.00006681072],"domain_scores_gemma":[0.9980742,0.00096937746,0.00023213551,0.00016889248,0.00048101475,0.00007437271],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031089108,0.0012667249,0.0018021574,0.006942181,0.00058237853,0.0023044925,0.0025097749,0.0013784262,0.0013936197],"category_scores_gemma":[0.004405786,0.00046128652,0.001578271,0.008170708,0.0011699508,0.0037749307,0.002528041,0.002277789,0.0010027904],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000059616494,0.00012654207,0.0006875539,0.009579909,0.00020726067,0.00021926276,0.00029451135,0.0015951575,0.00448365,0.041006498,0.024339614,0.91740036],"study_design_scores_gemma":[0.000039413124,0.00007158881,0.00271706,0.008724398,0.00053653406,0.0019496033,0.0004168986,0.008409753,0.005846691,0.07017088,0.9009839,0.00013323831],"about_ca_topic_score_codex":0.008624113,"about_ca_topic_score_gemma":0.008014082,"teacher_disagreement_score":0.008624113,"about_ca_system_score_codex":0.0026652252,"about_ca_system_score_gemma":0.005833511,"threshold_uncertainty_score":0.019337654},"labels":[],"label_agreement":null},{"id":"W2128365672","doi":"10.1186/2041-1480-4-18","title":"The mouse pathology ontology, MPATH; structure and applications","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"National Cancer Institute; National Human Genome Research Institute; National Institutes of Health; European Commission; Ellison Medical Foundation","keywords":"Computer science; Ontology; Data science; Information retrieval","score_opus":0.006309531117010112,"score_gpt":0.2474151125741754,"score_spread":0.2411055814571653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128365672","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01237049,0.0018465844,0.76788485,0.0036916356,0.00040538513,0.0016284058,0.10304009,0.07902811,0.030104602],"genre_scores_gemma":[0.050448097,0.0029609485,0.778851,0.0014467088,0.00019953931,0.0022503915,0.1467792,0.007061756,0.010002399],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.998793,0.00017432384,0.0002436832,0.0002458264,0.0004817548,0.000061311184],"domain_scores_gemma":[0.99765825,0.00082021573,0.0003700088,0.00052683655,0.0004467898,0.00017785822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032378493,0.0010456467,0.00039126203,0.0059600403,0.0008297375,0.0023162796,0.0024137625,0.0011394661,0.0062943264],"category_scores_gemma":[0.0061297957,0.0008187363,0.0014782266,0.0048216083,0.0011669602,0.005007366,0.0030411424,0.0017112356,0.0029700198],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000299307,0.00027526097,0.006729427,0.004732655,0.00020545268,0.0014186474,0.0017516051,0.013665246,0.02036238,0.27812496,0.28340924,0.38902578],"study_design_scores_gemma":[0.000043825825,0.000047135603,0.004040323,0.00066470844,0.000082425206,0.0015528798,0.0002323417,0.033302862,0.0075577633,0.07221794,0.8801642,0.00009369422],"about_ca_topic_score_codex":0.0069110803,"about_ca_topic_score_gemma":0.006379443,"teacher_disagreement_score":0.0069110803,"about_ca_system_score_codex":0.0025670792,"about_ca_system_score_gemma":0.0041932943,"threshold_uncertainty_score":0.021056652},"labels":[],"label_agreement":null},{"id":"W2128795341","doi":"10.1109/cbms.2009.5255451","title":"Scenario-oriented information extraction from electronic health records","year":2009,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Computer science; Matching (statistics); Information retrieval; Electronic health record; Task (project management); Set (abstract data type); Health records; Information extraction; Rank (graph theory); Data mining; Health care; Data science; Medicine","score_opus":0.007250516698653354,"score_gpt":0.27929599098494073,"score_spread":0.2720454742862874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128795341","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042511564,0.0016145498,0.932519,0.002038354,0.000118519674,0.0017298727,0.011686261,0.0036255175,0.0041564386],"genre_scores_gemma":[0.14818028,0.001327704,0.8323443,0.00020798395,0.00007040507,0.00054493436,0.016533995,0.00008882694,0.00070161826],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9963721,0.0012929937,0.0006236382,0.00053123233,0.001048173,0.00013182775],"domain_scores_gemma":[0.98881185,0.0077662542,0.0010208924,0.0008787182,0.0013037304,0.00021853801],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026482733,0.0012733653,0.00086353376,0.009959299,0.00103822,0.0025916141,0.0015298518,0.0015199148,0.002565911],"category_scores_gemma":[0.016955828,0.0006197397,0.0019031949,0.0077612894,0.0005491456,0.0035249926,0.0020771795,0.00113763,0.0012506447],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064022554,0.0007435655,0.023871487,0.004178591,0.00053274597,0.0065933107,0.0024575295,0.079717174,0.0205302,0.053221814,0.02679075,0.78072256],"study_design_scores_gemma":[0.00021693796,0.0003915567,0.016854662,0.0015368336,0.00064136094,0.00610982,0.0042722584,0.65557504,0.04068997,0.16822258,0.1051798,0.00030916897],"about_ca_topic_score_codex":0.0025877382,"about_ca_topic_score_gemma":0.004267237,"teacher_disagreement_score":0.009959299,"about_ca_system_score_codex":0.00091655256,"about_ca_system_score_gemma":0.0028269377,"threshold_uncertainty_score":0.014005601},"labels":[],"label_agreement":null},{"id":"W2128972928","doi":"10.1186/1745-6215-15-369","title":"Linked publications from a single trial: a thread of evidence","year":2014,"lang":"en","type":"editorial","venue":"Trials","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; University of Ottawa","funders":"Cancer Research UK","keywords":"Medicine; Clinical trial; Medical research; Alternative medicine; Medical journal; MEDLINE; Data science; Engineering ethics; Family medicine; Computer science; Pathology","score_opus":0.34439745857873927,"score_gpt":0.4496192824422373,"score_spread":0.10522182386349804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2128972928","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008056817,0.016165644,0.001160081,0.14594175,0.83533454,0.000057792433,0.00016110431,0.000126619,0.0009718979],"genre_scores_gemma":[0.0006995053,0.013571133,0.0011720215,0.08652247,0.8938453,0.00010460089,0.00008642962,0.00015196744,0.003846577],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96232367,0.010632291,0.008348743,0.0029911776,0.014976956,0.0007271453],"domain_scores_gemma":[0.7173694,0.21759921,0.014574854,0.0061273505,0.037753824,0.0065753898],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04627279,0.0027173918,0.0052254745,0.008612172,0.0037505787,0.0147117805,0.0051862476,0.019934721,0.007433525],"category_scores_gemma":[0.19219857,0.0021313182,0.0036775598,0.0056651705,0.0073842914,0.014082021,0.0051293275,0.037726715,0.005375928],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007083997,0.0000122499805,0.0000584289,0.001585998,0.00013955598,0.00015744483,0.000094803116,0.000040981162,0.000066995846,0.0015278088,0.98071057,0.015534216],"study_design_scores_gemma":[0.00011340948,0.000029070074,0.0002896761,0.005518907,0.0003021102,0.00044520645,0.00011295215,0.00024117007,0.00017378933,0.006922481,0.9857982,0.000053033695],"about_ca_topic_score_codex":0.0012877071,"about_ca_topic_score_gemma":0.0033779782,"teacher_disagreement_score":0.9537272,"about_ca_system_score_codex":0.0048544495,"about_ca_system_score_gemma":0.0076332144,"threshold_uncertainty_score":0.24471682},"labels":[],"label_agreement":null},{"id":"W2129500526","doi":"10.1016/j.cmpb.2011.01.002","title":"Systematized nomenclature of medicine clinical terms (SNOMED CT) to represent computed tomography procedures","year":2011,"lang":"en","type":"article","venue":"Computer Methods and Programs in Biomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":71,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Memorial University of Newfoundland; Newfoundland and Labrador Centre for Applied Health Research","funders":"Government of Canada","keywords":"SNOMED CT; Systematized Nomenclature of Medicine; Computed tomography; Nomenclature; Medicine; Radiology; Medical physics; Artificial intelligence; Computer science; Information retrieval; Terminology; Taxonomy (biology); Linguistics","score_opus":0.1158773027293621,"score_gpt":0.41601120764258165,"score_spread":0.30013390491321956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129500526","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02014162,0.0054011764,0.7774966,0.0022751852,0.0021269806,0.0024118333,0.14460656,0.006991835,0.03854822],"genre_scores_gemma":[0.049892846,0.0027879796,0.8463189,0.0014623348,0.00019225648,0.0012327477,0.09401774,0.00077192544,0.003323227],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99781656,0.0006166807,0.00074250094,0.00025350897,0.00048671765,0.00008406047],"domain_scores_gemma":[0.9964019,0.0016137799,0.00051047426,0.0004183418,0.00088284083,0.00017252733],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016631095,0.0010955031,0.0006276647,0.011353528,0.0012079375,0.0025510022,0.0015221556,0.0012535057,0.007096382],"category_scores_gemma":[0.0057441243,0.00039287383,0.001484444,0.012955993,0.0011972996,0.0021774834,0.0014307444,0.0017314645,0.002320581],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046585494,0.00030679457,0.007913845,0.0073505044,0.00043370394,0.0024655573,0.0030766842,0.020522693,0.02495519,0.49825162,0.14105302,0.29320452],"study_design_scores_gemma":[0.00009463519,0.0001228957,0.0048973802,0.0020335845,0.00040544392,0.0035413448,0.00074988103,0.021651333,0.008307582,0.09540736,0.862655,0.00013363517],"about_ca_topic_score_codex":0.0094235325,"about_ca_topic_score_gemma":0.015908968,"teacher_disagreement_score":0.011353528,"about_ca_system_score_codex":0.0018202784,"about_ca_system_score_gemma":0.0068714404,"threshold_uncertainty_score":0.023739755},"labels":[],"label_agreement":null},{"id":"W2129769245","doi":"","title":"Adapting a General Semantic Interpretation Approach to Biological Event Extraction","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Event (particle physics); Negation; Biomedical text mining; Natural language processing; Task (project management); Generalization; Focus (optics); Artificial intelligence; Embedding; Domain (mathematical analysis); Property (philosophy); Programming language; Text mining; Mathematics","score_opus":0.05839219796191509,"score_gpt":0.3001508366762506,"score_spread":0.2417586387143355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129769245","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017989065,0.00010048955,0.99317884,0.00043020755,0.00004723615,0.00013452285,0.0002708301,0.0027456782,0.0012934057],"genre_scores_gemma":[0.041388944,0.00025432656,0.9535492,0.0003757527,0.00011418235,0.00013848332,0.0015327515,0.00070660445,0.0019397666],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99498135,0.0016187779,0.0006651327,0.001284374,0.0012859827,0.00016440645],"domain_scores_gemma":[0.99348265,0.0022716764,0.0004038759,0.0022694771,0.0014277804,0.00014447064],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065386114,0.001953349,0.0013646577,0.0046827337,0.0012474153,0.0046724663,0.0029128056,0.0018815008,0.004670504],"category_scores_gemma":[0.010749407,0.0010907671,0.0035224846,0.0038497145,0.0028442645,0.008954269,0.0049042213,0.0038849267,0.0027146407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029112666,0.00031489017,0.004187769,0.0018215392,0.00042315078,0.0010963362,0.0036834518,0.019992525,0.043247994,0.30385748,0.022659281,0.5984244],"study_design_scores_gemma":[0.00008421963,0.00013974943,0.0021863116,0.00025106178,0.00039855394,0.0014415202,0.0011501688,0.34524718,0.047099873,0.46287936,0.13894884,0.0001731265],"about_ca_topic_score_codex":0.0020775897,"about_ca_topic_score_gemma":0.0034022555,"teacher_disagreement_score":0.0065386114,"about_ca_system_score_codex":0.001411385,"about_ca_system_score_gemma":0.0022238025,"threshold_uncertainty_score":0.034579933},"labels":[],"label_agreement":null},{"id":"W2129850776","doi":"10.1177/0162243906287545","title":"A New Clinical Collective for French Cancer Genetics","year":2006,"lang":"en","type":"article","venue":"Science Technology & Human Values","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":47,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Cancer genetics; Clinical trial; Cancer; Articulation (sociology); Breast cancer; Medicine; Genetics; Bioinformatics; Biology; Political science","score_opus":0.03633243536601799,"score_gpt":0.3898647637021895,"score_spread":0.3535323283361715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2129850776","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36096537,0.03825714,0.07354819,0.31045625,0.0064375233,0.0009705559,0.006532857,0.0012465875,0.2015855],"genre_scores_gemma":[0.86777884,0.0057497174,0.058274444,0.0091096675,0.0024455902,0.00060881855,0.004137038,0.0002463197,0.051649455],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99200094,0.004727165,0.0003495494,0.0009645195,0.0013653378,0.0005924453],"domain_scores_gemma":[0.96662897,0.016619084,0.0031973582,0.0028508685,0.006986031,0.0037177939],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0147770895,0.0004033796,0.00039687168,0.006907998,0.006280776,0.006956433,0.0006507319,0.0018233584,0.010157243],"category_scores_gemma":[0.023767704,0.00020868302,0.00049458956,0.008272725,0.0034723138,0.0038160272,0.003600543,0.0014782654,0.000933682],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025643708,0.00008151598,0.057250313,0.00057145447,0.00006606712,0.0024809644,0.089960076,0.0010312505,0.00324513,0.3296881,0.14941958,0.36594915],"study_design_scores_gemma":[0.000030138965,0.00006312806,0.028284755,0.00029856164,0.0000243676,0.0009411415,0.02586416,0.00044651362,0.0005226515,0.01861838,0.92485696,0.000049359383],"about_ca_topic_score_codex":0.030349622,"about_ca_topic_score_gemma":0.027614146,"teacher_disagreement_score":0.030349622,"about_ca_system_score_codex":0.0121039655,"about_ca_system_score_gemma":0.015890015,"threshold_uncertainty_score":0.08782077},"labels":[],"label_agreement":null},{"id":"W2130063192","doi":"10.1145/1854776.1854820","title":"Unsupervised mapping of sentences to biomedical concepts based on integrated information retrieval model and clustering","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Unified Medical Language System; Information retrieval; Cluster analysis; Annotation; Natural language processing; Task (project management); Scope (computer science); Artificial intelligence; Matching (statistics); Thesaurus; Precision and recall","score_opus":0.015959747692221745,"score_gpt":0.27503783881919563,"score_spread":0.2590780911269739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130063192","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05762292,0.0009543164,0.9316953,0.0005246334,0.00010413296,0.0007266363,0.0012368602,0.004903629,0.0022315255],"genre_scores_gemma":[0.2003228,0.00054454413,0.78974193,0.0002714911,0.0001012012,0.00072069,0.005017822,0.00033191728,0.0029475226],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971915,0.0007467532,0.0003229298,0.00089388766,0.0006817263,0.00016314579],"domain_scores_gemma":[0.99612683,0.0016270583,0.00033766663,0.00056330475,0.0012545766,0.00009048341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002231307,0.0013247889,0.0014176307,0.008683351,0.0008674105,0.001672462,0.0020163225,0.001458486,0.0017125639],"category_scores_gemma":[0.00837125,0.00043010197,0.0019215489,0.0060168733,0.0007811866,0.003324241,0.0013286714,0.0010546121,0.0018585507],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006414334,0.00073419715,0.008132645,0.0011884697,0.0004057755,0.0007839937,0.0023741042,0.046583466,0.07888933,0.017192943,0.019614825,0.82345873],"study_design_scores_gemma":[0.00010627273,0.00030128547,0.008626115,0.000094931966,0.00028473898,0.0009049008,0.0008281093,0.9106459,0.03696572,0.028274782,0.012813709,0.00015354497],"about_ca_topic_score_codex":0.0085436925,"about_ca_topic_score_gemma":0.007573678,"teacher_disagreement_score":0.008683351,"about_ca_system_score_codex":0.0014088879,"about_ca_system_score_gemma":0.00224939,"threshold_uncertainty_score":0.01698792},"labels":[],"label_agreement":null},{"id":"W2130549711","doi":"10.1093/bioinformatics/btr337","title":"Annotation concept synthesis and enrichment analysis: a logic-based approach to the interpretation of high-throughput experiments","year":2011,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Interpretation (philosophy); Annotation; Throughput; Computer science; Programming language; Artificial intelligence","score_opus":0.029208102939729797,"score_gpt":0.26855426814801336,"score_spread":0.23934616520828356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130549711","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031375445,0.00015612865,0.9948809,0.0002739111,0.000021417334,0.0001597132,0.00020678164,0.0005957086,0.00056780357],"genre_scores_gemma":[0.05806025,0.00021074228,0.9398205,0.0003061224,0.000057659734,0.00042212757,0.0005956658,0.00008678391,0.00044007492],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9896722,0.0043048677,0.0006462156,0.0016900084,0.0033417952,0.00034496066],"domain_scores_gemma":[0.96638507,0.026628673,0.0017580842,0.0017602483,0.0030601674,0.0004077687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014292174,0.0026031458,0.0014545323,0.0064649805,0.0013659456,0.0044668107,0.0038490614,0.001262193,0.0034645677],"category_scores_gemma":[0.024324542,0.00070522306,0.0037797454,0.0036106606,0.0035214636,0.0040008426,0.0028981378,0.0030749603,0.00077595754],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011617348,0.0008274314,0.008717907,0.0031990423,0.0010412447,0.0016316104,0.0017281964,0.18871275,0.039579567,0.24280253,0.005141851,0.50545615],"study_design_scores_gemma":[0.000105674015,0.00018796707,0.0011845588,0.0002776196,0.0003001051,0.00034473566,0.00026651012,0.6198054,0.028993236,0.34006312,0.008369319,0.00010172384],"about_ca_topic_score_codex":0.0027111715,"about_ca_topic_score_gemma":0.0022537229,"teacher_disagreement_score":0.014292174,"about_ca_system_score_codex":0.002817674,"about_ca_system_score_gemma":0.004319505,"threshold_uncertainty_score":0.07558519},"labels":[],"label_agreement":null},{"id":"W2130705682","doi":"10.1109/hicss.2005.121","title":"BiRD: A Strategy to Autonomously Supplement Clinical Practice Guidelines with Related Clinical Studies","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Information retrieval; XML; Query language; Query expansion; World Wide Web","score_opus":0.15950956064836178,"score_gpt":0.5120770212633639,"score_spread":0.3525674606150021,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130705682","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070735323,0.0005635735,0.9681966,0.002014034,0.00013772215,0.0007887804,0.0012432362,0.016482256,0.00350018],"genre_scores_gemma":[0.025522618,0.00018923258,0.9700702,0.00053360284,0.000053196625,0.00018289141,0.0014207491,0.0005694582,0.0014579804],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99195087,0.0042989785,0.0010346378,0.00089876103,0.0016155252,0.00020113945],"domain_scores_gemma":[0.96466154,0.022959784,0.0023705515,0.005494481,0.0036526492,0.00086107873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018323889,0.00095519156,0.0012983673,0.009522981,0.0008193565,0.004028793,0.00201655,0.0017296005,0.0056728874],"category_scores_gemma":[0.05004718,0.0008292801,0.0011285675,0.004739139,0.0011254804,0.0064197388,0.0049317284,0.0012657046,0.003472555],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010768011,0.00027685362,0.004248977,0.0018635677,0.00036697334,0.000907788,0.0033941786,0.0060545234,0.01691754,0.048553377,0.046991605,0.8693478],"study_design_scores_gemma":[0.000661239,0.0009498392,0.0042745373,0.0012180299,0.0006796224,0.002265331,0.0020482668,0.16178715,0.050625287,0.11315174,0.6618773,0.0004617548],"about_ca_topic_score_codex":0.0038340841,"about_ca_topic_score_gemma":0.005929076,"teacher_disagreement_score":0.018323889,"about_ca_system_score_codex":0.0011882122,"about_ca_system_score_gemma":0.0035583053,"threshold_uncertainty_score":0.09690714},"labels":[],"label_agreement":null},{"id":"W2130930065","doi":"10.1186/gb-2008-9-s2-s7","title":"Text mining for biology - the way forward: opinions from leading scientists","year":2008,"lang":"en","type":"article","venue":"Genome biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":101,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"European Social Fund; European Science Foundation","keywords":"Workflow; Computer science; Data science; Process (computing); Annotation; World Wide Web; Artificial intelligence","score_opus":0.03921410073481182,"score_gpt":0.3064619274795019,"score_spread":0.2672478267446901,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2130930065","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016650035,0.09522696,0.004780393,0.88092554,0.01341536,0.000050800103,0.0001785982,0.00006528762,0.0036920612],"genre_scores_gemma":[0.04799176,0.325973,0.03178942,0.54215175,0.038176462,0.00023286062,0.0008695845,0.00050933106,0.012305858],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.9697188,0.010405647,0.003843957,0.0020546855,0.012749508,0.0012274401],"domain_scores_gemma":[0.78567404,0.11125568,0.0049044527,0.0038017808,0.08060794,0.013756053],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.058405917,0.0010286572,0.0016013492,0.003872145,0.0036325732,0.013759833,0.002074788,0.0077511678,0.0028706098],"category_scores_gemma":[0.09072199,0.0004076924,0.0012227055,0.007212819,0.0050907223,0.013354809,0.0055599795,0.014512975,0.0020692616],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016048974,0.00006246244,0.0024998405,0.0021967152,0.0001198311,0.00044235328,0.010150148,0.0005138258,0.0019533508,0.009878438,0.60487705,0.36714542],"study_design_scores_gemma":[0.000034042914,0.0000494166,0.0017922205,0.005469755,0.000103211416,0.00040109546,0.015856296,0.0006482449,0.0010157871,0.023686444,0.9508537,0.0000898982],"about_ca_topic_score_codex":0.0051738117,"about_ca_topic_score_gemma":0.008256559,"teacher_disagreement_score":0.94159406,"about_ca_system_score_codex":0.004670814,"about_ca_system_score_gemma":0.013194909,"threshold_uncertainty_score":0.30888373},"labels":[],"label_agreement":null},{"id":"W2131454999","doi":"10.12927/hcq..18096","title":"CIHR Research: Translating a Broad Term into Real-World Applications: CIHR's Successful Approach to Knowledge Translation","year":2006,"lang":"en","type":"article","venue":"Healthcare Quarterly","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institutes of Health Research","funders":"","keywords":"Best practice; Term (time); Knowledge translation; Computer science; Political science; Knowledge management; Physics","score_opus":0.08903968980790332,"score_gpt":0.3985398877615011,"score_spread":0.3095001979535978,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131454999","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077138087,0.0051739137,0.83339024,0.054468874,0.0032021794,0.00088923814,0.0024271093,0.0073812655,0.085353345],"genre_scores_gemma":[0.0795259,0.0049604797,0.8823966,0.006254094,0.0012056916,0.00082030037,0.0039043438,0.002814102,0.0181185],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9649988,0.019457364,0.0035892623,0.0028652248,0.008168745,0.00092062284],"domain_scores_gemma":[0.8501527,0.089529194,0.003961782,0.02988792,0.024324106,0.0021443113],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04621121,0.0013186327,0.0015080371,0.015081352,0.0043130103,0.014649739,0.0040361257,0.0051113158,0.010759968],"category_scores_gemma":[0.11628858,0.0012168649,0.0014197222,0.01668277,0.009249664,0.028587235,0.011076198,0.008835755,0.0066270726],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010249581,0.00019512266,0.00054912287,0.002207969,0.00007678727,0.00036715684,0.0098993005,0.0012643654,0.00557795,0.54376704,0.08125962,0.35473308],"study_design_scores_gemma":[0.000081347374,0.000098058554,0.0011847791,0.0012615867,0.00010325823,0.0010004356,0.004687828,0.009595613,0.013031023,0.32302392,0.64574116,0.00019092593],"about_ca_topic_score_codex":0.008984983,"about_ca_topic_score_gemma":0.0063242363,"teacher_disagreement_score":0.9937358,"about_ca_system_score_codex":0.006264239,"about_ca_system_score_gemma":0.019803949,"threshold_uncertainty_score":0.2443912},"labels":[],"label_agreement":null},{"id":"W2131660156","doi":"10.1093/nar/gkp440","title":"BioPortal: ontologies and integrated data resources at the click of a mouse","year":2009,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":884,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"National Human Genome Research Institute","keywords":"Biology; Computational biology; Genetics; Bioinformatics","score_opus":0.09505631666808116,"score_gpt":0.3857343446169439,"score_spread":0.2906780279488627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2131660156","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031670572,0.0033326019,0.49371675,0.009528632,0.0024597156,0.0011974442,0.13870063,0.27857065,0.0693265],"genre_scores_gemma":[0.023653869,0.005587602,0.48946428,0.008325702,0.0013507729,0.002075574,0.36749303,0.048808016,0.05324126],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975178,0.00063868554,0.00038490276,0.0003287747,0.0008916192,0.00023824298],"domain_scores_gemma":[0.99455345,0.0017883232,0.00046452307,0.0016093988,0.0006294708,0.0009548225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005307113,0.0016985987,0.0015592485,0.0066799014,0.0012739683,0.005235287,0.0027520563,0.003504436,0.08149143],"category_scores_gemma":[0.0117877815,0.0016677426,0.0013625382,0.0064180484,0.0012785711,0.01247326,0.010699465,0.0036898523,0.07540285],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004508236,0.00012717363,0.000762571,0.0013385133,0.000098805926,0.00055719976,0.00043615222,0.00031954754,0.0063185757,0.031952966,0.8540852,0.10355237],"study_design_scores_gemma":[0.00013183603,0.00003205847,0.0010566701,0.00048151237,0.000041458028,0.00055862754,0.0001390717,0.0013211425,0.0032074312,0.04235757,0.9505871,0.00008554878],"about_ca_topic_score_codex":0.001899342,"about_ca_topic_score_gemma":0.00287355,"teacher_disagreement_score":0.08149143,"about_ca_system_score_codex":0.0010323446,"about_ca_system_score_gemma":0.0023979486,"threshold_uncertainty_score":0.27261603},"labels":[],"label_agreement":null},{"id":"W2132047973","doi":"10.1186/1749-8546-5-43","title":"Integrating findings of traditional medicine with modern pharmaceutical research: the potential role of linked open data","year":2010,"lang":"en","type":"editorial","venue":"Chinese Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Engineering and Physical Sciences Research Council; Science Foundation Ireland","keywords":"Identification (biology); Data science; Data integration; Traditional Chinese medicine; Computer science; Medicine; Alternative medicine; Knowledge management; Traditional medicine; Management science; Engineering; Data mining; Biology","score_opus":0.11950800426458337,"score_gpt":0.4281862157215826,"score_spread":0.30867821145699925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2132047973","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00013064731,0.014984318,0.0030226288,0.106095254,0.8740569,0.00003871382,0.000061324514,0.00012177146,0.0014884961],"genre_scores_gemma":[0.001498111,0.018129557,0.0028038195,0.039247684,0.9330598,0.000054544445,0.00005167697,0.00011305279,0.005041752],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9898493,0.0034785403,0.0015578506,0.0008314743,0.004056193,0.00022661978],"domain_scores_gemma":[0.88949335,0.08537451,0.003180618,0.0020862664,0.016270164,0.0035950425],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.016353846,0.0018712851,0.0019057635,0.004745709,0.0026577017,0.011612594,0.0033144082,0.011432914,0.0027623682],"category_scores_gemma":[0.049433693,0.0010305552,0.00205726,0.0031538354,0.0057071196,0.009990812,0.00280466,0.025350006,0.0019527291],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044678363,0.00002491487,0.00011642012,0.00085917674,0.0000937786,0.00047488595,0.0002676635,0.00018535416,0.00023961371,0.0043257405,0.96433616,0.029031629],"study_design_scores_gemma":[0.000032279844,0.000019929603,0.00026169437,0.00088774617,0.00008185487,0.0005214721,0.0001876264,0.0006666076,0.0001905551,0.0057937806,0.9913202,0.000036321493],"about_ca_topic_score_codex":0.0022717004,"about_ca_topic_score_gemma":0.0052808253,"teacher_disagreement_score":0.99668556,"about_ca_system_score_codex":0.0029642433,"about_ca_system_score_gemma":0.003908098,"threshold_uncertainty_score":0.086488426},"labels":[],"label_agreement":null},{"id":"W2133018204","doi":"10.1186/1471-2164-11-s4-s24","title":"Algorithms and semantic infrastructure for mutation impact extraction and grounding","year":2010,"lang":"en","type":"article","venue":"BMC Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; University of New Brunswick","funders":"Natural Sciences and Engineering Research Council of Canada; New Brunswick Innovation Foundation","keywords":"Computer science; Precision and recall; Mutation; Relationship extraction; Information retrieval; Ontology; Information extraction; UniProt; Computational biology; World Wide Web; Biology; Genetics; Gene","score_opus":0.013982663067837145,"score_gpt":0.3009710052028088,"score_spread":0.28698834213497165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133018204","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056756297,0.0006085079,0.94444484,0.00086801813,0.00009436074,0.00050279096,0.0077680047,0.035942174,0.004095622],"genre_scores_gemma":[0.04770049,0.00080744846,0.9187123,0.0003989467,0.000085350635,0.00046430298,0.028852135,0.001458992,0.0015199386],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953277,0.0007317388,0.00085861,0.0011578209,0.0017125043,0.0002116669],"domain_scores_gemma":[0.9925055,0.003325388,0.00081762695,0.0015598233,0.0015930459,0.00019857234],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004370386,0.0017450966,0.001141196,0.011192534,0.0014473674,0.0049567954,0.0025740983,0.0022170024,0.0075397547],"category_scores_gemma":[0.013618119,0.00073731405,0.003621095,0.006879427,0.001408464,0.0055887247,0.004202531,0.0021071236,0.0055044745],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035730985,0.00062870973,0.009914815,0.0035652053,0.000547622,0.0015321047,0.0012415707,0.03427714,0.02297983,0.090233356,0.053191934,0.7815304],"study_design_scores_gemma":[0.00013807885,0.00014365338,0.005502122,0.0009652561,0.00037898478,0.0016006362,0.00081784395,0.40628576,0.050673112,0.2918157,0.24150786,0.00017109344],"about_ca_topic_score_codex":0.004391465,"about_ca_topic_score_gemma":0.005329707,"teacher_disagreement_score":0.011192534,"about_ca_system_score_codex":0.0021271482,"about_ca_system_score_gemma":0.0045510926,"threshold_uncertainty_score":0.025222957},"labels":[],"label_agreement":null},{"id":"W2133525847","doi":"10.1093/bioinformatics/btl235","title":"Integrating image data into biomedical text categorization","year":2006,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":118,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada; Xerox Foundation","keywords":"Computer science; Annotation; Categorization; Information retrieval; Task (project management); Text categorization; Feature (linguistics); Artificial intelligence; Natural language processing","score_opus":0.015740115506897037,"score_gpt":0.27755073193468244,"score_spread":0.2618106164277854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133525847","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10884064,0.006781992,0.84506637,0.0029286987,0.0005825154,0.000953729,0.009363782,0.01357989,0.011902402],"genre_scores_gemma":[0.16869038,0.0019982695,0.8128728,0.00040445698,0.00048508693,0.00033421352,0.011146557,0.0003189728,0.0037492805],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990363,0.00016104116,0.00010479038,0.00025203679,0.00036534684,0.00008048799],"domain_scores_gemma":[0.9963206,0.0014579683,0.00031800507,0.0006578663,0.0010842286,0.00016135351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011219992,0.00093128055,0.0009563782,0.009778846,0.0005494352,0.0022753077,0.001083211,0.001498062,0.0031567556],"category_scores_gemma":[0.0054279445,0.00033390019,0.00092079956,0.005108669,0.00067501154,0.0030697025,0.001324425,0.00091150875,0.0030247173],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028680908,0.00040307594,0.0060221893,0.00072129176,0.00010846376,0.0003440718,0.00019563979,0.00362822,0.093587935,0.0035191628,0.011696801,0.8794864],"study_design_scores_gemma":[0.00017079165,0.0009899054,0.053649984,0.00051828666,0.0006003064,0.0026293898,0.0015677648,0.39424908,0.32659242,0.09884266,0.1198579,0.00033149472],"about_ca_topic_score_codex":0.00205477,"about_ca_topic_score_gemma":0.0027658117,"teacher_disagreement_score":0.009778846,"about_ca_system_score_codex":0.0005417443,"about_ca_system_score_gemma":0.0005725572,"threshold_uncertainty_score":0.010560393},"labels":[],"label_agreement":null},{"id":"W2133550211","doi":"10.1093/bioinformatics/btr058","title":"A common layer of interoperability for biomedical ontologies based on OWL EL","year":2011,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"National Human Genome Research Institute; Biotechnology and Biological Sciences Research Council; European Commission; Natural Sciences and Engineering Research Council of Canada; National Institutes of Health; European Bioinformatics Institute","keywords":"Computer science; Interoperability; Ontology; Open Biomedical Ontologies; Web Ontology Language; IDEF5; Semantics (computer science); Knowledge representation and reasoning; Ontology components; Upper ontology; Description logic; Software; Information retrieval; Semantic Web; Software engineering; World Wide Web; Programming language; Artificial intelligence; Ontology alignment","score_opus":0.057446294203414114,"score_gpt":0.30675205372400544,"score_spread":0.2493057595205913,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133550211","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008901391,0.00018316887,0.9679,0.0010196663,0.00010845165,0.00044482073,0.0010966402,0.012653535,0.0076924236],"genre_scores_gemma":[0.10507289,0.00026415143,0.882709,0.0007974674,0.000067759975,0.00051459676,0.0052888207,0.0016064357,0.003678855],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9899768,0.0028571754,0.0016661291,0.0011734875,0.003747214,0.0005791947],"domain_scores_gemma":[0.98283374,0.004107384,0.0012809249,0.00843246,0.0027932154,0.00055229454],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012057998,0.00085046026,0.0010164928,0.0028144857,0.0017316607,0.0069727106,0.0035462186,0.001769876,0.004038912],"category_scores_gemma":[0.025064815,0.0012012081,0.0030847038,0.002545581,0.0024118815,0.00940431,0.007785522,0.0042343494,0.0022202653],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007052248,0.0007761928,0.005324319,0.0010849007,0.00052348466,0.0017860516,0.0042556343,0.02643102,0.031701475,0.5418651,0.03558814,0.34995848],"study_design_scores_gemma":[0.00015819514,0.00021335408,0.0032946651,0.00075911224,0.00037491904,0.0016555097,0.0009466902,0.24097832,0.054735668,0.387515,0.30915844,0.00021013748],"about_ca_topic_score_codex":0.006744383,"about_ca_topic_score_gemma":0.0048336126,"teacher_disagreement_score":0.012057998,"about_ca_system_score_codex":0.0018468738,"about_ca_system_score_gemma":0.00468078,"threshold_uncertainty_score":0.06376958},"labels":[],"label_agreement":null},{"id":"W2133601033","doi":"10.1093/bioinformatics/btp602","title":"Evaluation of linguistic features useful in extraction of interactions from PubMed; Application to annotating known, high-throughput and predicted interactions in I2D","year":2009,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":77,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"U.S. National Library of Medicine; Ontario Genomics; Ontario Genomics Institute; Genome Canada","keywords":"Computer science; Sentence; Identification (biology); Recall; Precision and recall; Software; Natural language processing; Information retrieval; Process (computing); Artificial intelligence; Data mining; Machine learning; Programming language; Biology","score_opus":0.02690545087677414,"score_gpt":0.33233209653783674,"score_spread":0.3054266456610626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133601033","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8493561,0.0052263965,0.07697894,0.0020166857,0.00021675602,0.0015794833,0.04192663,0.015820256,0.0068787336],"genre_scores_gemma":[0.56926537,0.0010610692,0.3923007,0.0003391572,0.00010223134,0.0007924326,0.03459542,0.00029715532,0.0012464212],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.996698,0.0012184442,0.0008625767,0.00052888534,0.0005683697,0.00012361039],"domain_scores_gemma":[0.9751209,0.01913846,0.001765548,0.0007416039,0.0026775203,0.0005560324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059854006,0.001101822,0.0008591355,0.009662319,0.0010664073,0.0016143203,0.0011033688,0.001078992,0.0018751792],"category_scores_gemma":[0.018291546,0.00026435906,0.00071305357,0.005032147,0.00038706244,0.0015085334,0.0012107014,0.00046340618,0.0007690819],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005479754,0.0021151102,0.108122535,0.010525835,0.00093403563,0.003622109,0.0027997443,0.009720945,0.1188782,0.0017934073,0.041388907,0.6946194],"study_design_scores_gemma":[0.0018667963,0.0029654664,0.28063488,0.0010558692,0.0022572582,0.007181465,0.004946539,0.4013218,0.22121535,0.0047447262,0.07128467,0.00052522606],"about_ca_topic_score_codex":0.00444016,"about_ca_topic_score_gemma":0.0085008275,"teacher_disagreement_score":0.009662319,"about_ca_system_score_codex":0.0008938867,"about_ca_system_score_gemma":0.0014336987,"threshold_uncertainty_score":0.03165418},"labels":[],"label_agreement":null},{"id":"W2133615853","doi":"10.3115/1567619.1567623","title":"Term generalization and synonym resolution for biological abstracts","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta","keywords":"Computer science; Synonym (taxonomy); Ontology; Information retrieval; Task (project management); Hierarchy; Generalization; Categorization; Gene ontology; WordNet; Text categorization; Field (mathematics); Function (biology); Artificial intelligence; Natural language processing; Gene; Biology; Mathematics; Genus","score_opus":0.020299980912479348,"score_gpt":0.2687625671374854,"score_spread":0.24846258622500608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2133615853","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052209392,0.0026344436,0.9299484,0.0012786258,0.00051038567,0.00077638857,0.0034603225,0.005756102,0.0034259288],"genre_scores_gemma":[0.19689113,0.0013814708,0.78499335,0.0003992149,0.00075602543,0.0008572783,0.010661404,0.0005793338,0.0034807364],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9894488,0.0026282216,0.001867891,0.0026582677,0.0029656382,0.00043121082],"domain_scores_gemma":[0.9735859,0.01575505,0.0023058073,0.0042128433,0.0036497456,0.00049067725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011503149,0.0015998618,0.0023158533,0.017847953,0.002567377,0.0027071766,0.002909261,0.0021599785,0.004550777],"category_scores_gemma":[0.043097693,0.0006344271,0.003103132,0.011944114,0.001362385,0.008364159,0.0038154274,0.0031333028,0.0034310694],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005296183,0.00022869333,0.005963189,0.0010516493,0.00027664722,0.00083563646,0.0014666918,0.011494293,0.019874414,0.02026569,0.025632007,0.91238153],"study_design_scores_gemma":[0.00027246794,0.00044159067,0.01447331,0.00047014342,0.00064754125,0.0050708232,0.0020930795,0.6366968,0.03794525,0.23509113,0.066442594,0.00035523163],"about_ca_topic_score_codex":0.0030922072,"about_ca_topic_score_gemma":0.0029078186,"teacher_disagreement_score":0.017847953,"about_ca_system_score_codex":0.0014517425,"about_ca_system_score_gemma":0.0023894792,"threshold_uncertainty_score":0.060835183},"labels":[],"label_agreement":null},{"id":"W2134125816","doi":"10.3115/1572392.1572399","title":"An unsupervised method for extracting domain-specific affixes in biological literature","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Artificial intelligence; Domain (mathematical analysis); Matching (statistics); Prefix; Set (abstract data type); Annotation; Pattern recognition (psychology); Feature extraction; Tree (set theory); Process (computing); Affix; Mathematics","score_opus":0.03210451322672726,"score_gpt":0.35362785944363695,"score_spread":0.3215233462169097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134125816","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011715372,0.00065651373,0.9754803,0.00025541463,0.000120578465,0.00031076235,0.0018721432,0.006963396,0.0026255306],"genre_scores_gemma":[0.022753011,0.0002873504,0.9712903,0.00007016519,0.0000744251,0.00029621136,0.0033915844,0.00023424473,0.0016026441],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984309,0.00027997277,0.00023920747,0.00053241156,0.00045838064,0.00005911942],"domain_scores_gemma":[0.99546534,0.0017638818,0.0005782935,0.00063893426,0.0014400391,0.0001135065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013687446,0.00093931315,0.000709852,0.010122582,0.0010520553,0.0012156762,0.0012702582,0.00079976424,0.0033947453],"category_scores_gemma":[0.006135716,0.0004974351,0.00088485464,0.006669162,0.00074747036,0.0020773683,0.001113217,0.0010445929,0.0040643374],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000074139236,0.0001344973,0.0040675853,0.0009046551,0.00015571144,0.00040383646,0.000630457,0.0016737853,0.07809764,0.0066292933,0.01306708,0.89416134],"study_design_scores_gemma":[0.00019691669,0.0006182502,0.040870983,0.00047567047,0.00059696846,0.01166641,0.0017223776,0.3245037,0.2650746,0.048585612,0.30519614,0.00049243233],"about_ca_topic_score_codex":0.0013320543,"about_ca_topic_score_gemma":0.0039218604,"teacher_disagreement_score":0.010122582,"about_ca_system_score_codex":0.00038982436,"about_ca_system_score_gemma":0.002232061,"threshold_uncertainty_score":0.011356592},"labels":[],"label_agreement":null},{"id":"W2134135406","doi":"10.3115/1614108.1614149","title":"Simultaneous identification of biomedical named-entity and functional relations using statistical parsing techniques","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Parsing; Relationship extraction; Artificial intelligence; Natural language processing; Identification (biology); Relation (database); Task (project management); Information extraction; Domain (mathematical analysis); Key (lock); Information retrieval; Data mining","score_opus":0.018468907424254492,"score_gpt":0.30818772643989417,"score_spread":0.2897188190156397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2134135406","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01593097,0.00041054373,0.97452366,0.000699621,0.000046885852,0.00011454234,0.001356944,0.0058587315,0.0010581167],"genre_scores_gemma":[0.114510454,0.0004525127,0.87478536,0.00032781425,0.00010509252,0.00035361375,0.0076929834,0.0009688466,0.0008032898],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993328,0.0032113015,0.0007338397,0.0012555971,0.0012935172,0.00017776254],"domain_scores_gemma":[0.9483016,0.041334596,0.0027842056,0.0037201985,0.0036331885,0.0002261325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0087976055,0.0015016414,0.0015003469,0.00550259,0.001276795,0.0026317053,0.0018780796,0.0017802854,0.0020341026],"category_scores_gemma":[0.025872737,0.0010784705,0.00203304,0.0044281804,0.0012701929,0.0052295043,0.001913281,0.0024331606,0.0025130205],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046580387,0.0005549243,0.018213382,0.0017066303,0.00053538696,0.0018533278,0.0019568312,0.04703929,0.12871529,0.051477928,0.01768327,0.7297979],"study_design_scores_gemma":[0.000085021304,0.00016921946,0.012317872,0.00024630936,0.00045361504,0.0015490636,0.0006877528,0.6276803,0.16608642,0.14760716,0.042852987,0.00026420728],"about_ca_topic_score_codex":0.0014254778,"about_ca_topic_score_gemma":0.0021287662,"teacher_disagreement_score":0.0087976055,"about_ca_system_score_codex":0.00087219465,"about_ca_system_score_gemma":0.0028762615,"threshold_uncertainty_score":0.04652673},"labels":[],"label_agreement":null},{"id":"W2135479399","doi":"10.1186/2041-1480-4-9","title":"Semantic querying of relational data for clinical intelligence: a semantic web services-based approach","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; McGill University; University of New Brunswick","funders":"","keywords":"Computer science; Semantic analytics; Semantic Web; Information retrieval; Semantic Web Stack; Semantic computing; Social Semantic Web; Semantic grid; World Wide Web; Relational database; Natural language processing","score_opus":0.09271358477013052,"score_gpt":0.37172928567003427,"score_spread":0.27901570089990374,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135479399","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051143705,0.0007481394,0.97988015,0.0048539,0.00009855499,0.0002655119,0.0004465083,0.0015879855,0.0070049507],"genre_scores_gemma":[0.12761475,0.001709701,0.86518145,0.001172734,0.00018428055,0.00024243467,0.0018040821,0.00033677727,0.0017538753],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.993515,0.0023311444,0.0008191645,0.00055891793,0.002524261,0.00025144446],"domain_scores_gemma":[0.9928398,0.0035763946,0.00046714695,0.001390426,0.0013217413,0.00040435445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011224819,0.00057616376,0.00089583365,0.003918042,0.0013785817,0.00786644,0.003102371,0.0018611585,0.0019541227],"category_scores_gemma":[0.010110883,0.0005532839,0.0019007799,0.005396689,0.0034655735,0.00942791,0.0040812106,0.0021974663,0.00085460127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002173062,0.00026944347,0.0025934142,0.0009497666,0.00023038448,0.0011923243,0.0033352044,0.023749642,0.011396368,0.7834874,0.016144192,0.15643455],"study_design_scores_gemma":[0.0000737655,0.00006858411,0.00088453194,0.0003574013,0.00016268893,0.001230174,0.0019645002,0.2872812,0.014083734,0.53381544,0.15997373,0.00010435124],"about_ca_topic_score_codex":0.0054638926,"about_ca_topic_score_gemma":0.004294832,"teacher_disagreement_score":0.011224819,"about_ca_system_score_codex":0.0029013269,"about_ca_system_score_gemma":0.0035724503,"threshold_uncertainty_score":0.059363246},"labels":[],"label_agreement":null},{"id":"W2135536555","doi":"10.1136/amiajnl-2012-001075","title":"MEDLINE clinical queries are robust when searching in recent publishing years","year":2012,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":54,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hamilton Health Sciences; Western University; McMaster University","funders":"Canadian Institutes of Health Research","keywords":"Publishing; Computer science; MEDLINE; Information retrieval; Data science; World Wide Web; Political science","score_opus":0.037338715332168856,"score_gpt":0.3358608361451467,"score_spread":0.29852212081297785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135536555","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2909411,0.27738246,0.067685604,0.057507392,0.0071213767,0.02087024,0.16085255,0.006960339,0.11067901],"genre_scores_gemma":[0.6475356,0.074062556,0.13160819,0.020710023,0.0068295966,0.0137824165,0.09779893,0.0021089127,0.0055638035],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.65150666,0.104211375,0.16147627,0.016343981,0.0631271,0.00333467],"domain_scores_gemma":[0.18756579,0.57273984,0.11744804,0.043498527,0.07574671,0.003001162],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.19002537,0.0013916866,0.004516065,0.063513495,0.0016785398,0.017866543,0.00405662,0.004032583,0.008193596],"category_scores_gemma":[0.715902,0.0016081089,0.0027608513,0.069403015,0.0028768675,0.020071957,0.006579086,0.0015740418,0.0057723653],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030949835,0.0002843847,0.19238861,0.09678664,0.00376395,0.0019247197,0.01206114,0.0012225969,0.0057606115,0.007929099,0.13851503,0.5362683],"study_design_scores_gemma":[0.0010329823,0.0013728,0.4166438,0.07729415,0.008254598,0.007985599,0.010828896,0.0041576833,0.007560257,0.02116663,0.4427381,0.0009645985],"about_ca_topic_score_codex":0.0028735048,"about_ca_topic_score_gemma":0.006112793,"teacher_disagreement_score":0.8099746,"about_ca_system_score_codex":0.0030840323,"about_ca_system_score_gemma":0.008690496,"threshold_uncertainty_score":0.9988429},"labels":[],"label_agreement":null},{"id":"W2135868586","doi":"10.1007/s10115-010-0341-9","title":"Classifier-based acronym extraction for business documents","year":2010,"lang":"en","type":"article","venue":"Knowledge and Information Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Acronym; Computer science; Classifier (UML); Information extraction; Precision and recall; Natural language processing; Artificial intelligence; Information retrieval; Data mining; Linguistics","score_opus":0.013582429754063777,"score_gpt":0.29530810560313925,"score_spread":0.2817256758490755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2135868586","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11138631,0.0048693847,0.75383025,0.002072001,0.0019033055,0.0017674795,0.05588575,0.0478805,0.020404976],"genre_scores_gemma":[0.14836518,0.0013536694,0.7916337,0.00026048018,0.00024883205,0.00043646112,0.050631378,0.0010667229,0.006003579],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99801123,0.00028389477,0.00037335278,0.0005406141,0.0006170304,0.0001738467],"domain_scores_gemma":[0.9952172,0.0015916397,0.0005214002,0.0005270334,0.0019381257,0.00020461307],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011063088,0.0011919582,0.0013460111,0.011873848,0.0017884474,0.002769464,0.0010996065,0.0012981837,0.006391279],"category_scores_gemma":[0.007566184,0.00041986757,0.0014548531,0.008095327,0.00035924432,0.0032276472,0.0012275116,0.0013462142,0.0072675343],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005388529,0.00033585436,0.007283065,0.0015236697,0.00020277115,0.0011515926,0.0006089104,0.0027733892,0.055720445,0.0092993975,0.08014411,0.8404179],"study_design_scores_gemma":[0.00033187077,0.00044870627,0.01951351,0.0010312939,0.0010136769,0.0048873974,0.002360484,0.31936187,0.17369176,0.039212905,0.43781784,0.0003286669],"about_ca_topic_score_codex":0.005980211,"about_ca_topic_score_gemma":0.008971906,"teacher_disagreement_score":0.011873848,"about_ca_system_score_codex":0.0011529444,"about_ca_system_score_gemma":0.0034596585,"threshold_uncertainty_score":0.021380961},"labels":[],"label_agreement":null},{"id":"W2136265964","doi":"10.1093/bioinformatics/btt613","title":"Mouse model phenotypes provide information about human drug targets","year":2013,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Drug repositioning; Phenotype; Computational biology; Drug; Identification (biology); Similarity (geometry); Drug development; Drug discovery; Biology; Phenotypic screening; Repurposing; Computer science; Bioinformatics; Genetics; Pharmacology; Artificial intelligence; Gene","score_opus":0.010892537661136478,"score_gpt":0.24047240124778707,"score_spread":0.22957986358665058,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136265964","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039561134,0.0031560871,0.26067916,0.0034572892,0.00016372619,0.00059102016,0.6182119,0.050079953,0.024099655],"genre_scores_gemma":[0.20537248,0.005850348,0.27834496,0.0017370422,0.000207575,0.0011495829,0.49585822,0.0063005704,0.00517924],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99843735,0.00039108368,0.00024379372,0.0003503063,0.0005283123,0.000049106686],"domain_scores_gemma":[0.9889031,0.0061430554,0.0016440115,0.0019547227,0.0009799794,0.00037503408],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028329513,0.001700542,0.0009246326,0.005006981,0.00034655991,0.0013435956,0.0016593949,0.0010476122,0.018533397],"category_scores_gemma":[0.013982685,0.0005717518,0.00096865447,0.0048229424,0.0004931373,0.0024270813,0.0017138877,0.0010875122,0.008899318],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025570479,0.00083093444,0.07391763,0.009130728,0.0008682726,0.0020140156,0.00069943175,0.023827381,0.054499194,0.047509477,0.35210028,0.43204567],"study_design_scores_gemma":[0.0006044877,0.0005851174,0.08817182,0.0019092538,0.0012175185,0.005553112,0.00041117377,0.050856233,0.057897642,0.12926872,0.6631454,0.00037951383],"about_ca_topic_score_codex":0.0015508182,"about_ca_topic_score_gemma":0.0018617361,"teacher_disagreement_score":0.018533397,"about_ca_system_score_codex":0.00052935863,"about_ca_system_score_gemma":0.0010988032,"threshold_uncertainty_score":0.062000453},"labels":[],"label_agreement":null},{"id":"W2136291980","doi":"10.1371/journal.pone.0089606","title":"Semantics in Support of Biodiversity Knowledge Discovery: An Introduction to the Biological Collections Ontology and Related Ontologies","year":2014,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":151,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"National Human Genome Research Institute; European Commission; National Institutes of Health; National Science Foundation","keywords":"Ontology; Open Biomedical Ontologies; Biodiversity; Computer science; Data science; Metadata; Upper ontology; Terminology; Information retrieval; Semantic Web; Ecology; World Wide Web; Biology; Ontology alignment","score_opus":0.036283573067702246,"score_gpt":0.24823328292966412,"score_spread":0.21194970986196188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136291980","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013644231,0.010329182,0.96011895,0.011226008,0.0009443735,0.00024276496,0.00077380514,0.0008153059,0.014185191],"genre_scores_gemma":[0.033474423,0.021039566,0.9297441,0.004173739,0.0025522306,0.0010281581,0.0022396233,0.00062098337,0.0051271184],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9890793,0.0040785996,0.0021566295,0.0010757382,0.0030522025,0.0005575848],"domain_scores_gemma":[0.98558414,0.009570183,0.00081184076,0.0018753957,0.001567826,0.0005906166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014467611,0.0014122268,0.002074631,0.009757353,0.003906608,0.012492523,0.0054186904,0.004782182,0.004319308],"category_scores_gemma":[0.020435045,0.0018029183,0.004381355,0.014556803,0.0083118435,0.034992736,0.009931841,0.010052388,0.001796109],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002135907,0.000052635616,0.00034697822,0.00042850545,0.000043611828,0.00029598913,0.0016081642,0.0026924736,0.00047904367,0.9303974,0.011336186,0.052297518],"study_design_scores_gemma":[0.0000132813575,0.000015042647,0.00025399507,0.0006992406,0.000040533774,0.0004359758,0.000593145,0.012073243,0.00048338954,0.70298946,0.28234276,0.00005991212],"about_ca_topic_score_codex":0.011814451,"about_ca_topic_score_gemma":0.0082869725,"teacher_disagreement_score":0.014467611,"about_ca_system_score_codex":0.0049798023,"about_ca_system_score_gemma":0.005085031,"threshold_uncertainty_score":0.07651293},"labels":[],"label_agreement":null},{"id":"W2136317480","doi":"10.1186/1752-0509-5-124","title":"Integrating systems biology models and biomedical ontologies","year":2011,"lang":"en","type":"article","venue":"BMC Systems Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"Biotechnology and Biological Sciences Research Council; Natural Sciences and Engineering Research Council of Canada; European Commission","keywords":"SBML; Systems biology; Computer science; Markup language; Modelling biological systems; Systems medicine; Software; Granularity; Biological network; Semantics (computer science); Ontology; Synthetic biology; Open Biomedical Ontologies; Data science; Theoretical computer science; Computational biology; XML; Biology; Semantic Web; Programming language; Artificial intelligence; World Wide Web; Ontology-based data integration","score_opus":0.08872993251183821,"score_gpt":0.30083478766305083,"score_spread":0.21210485515121263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136317480","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039602127,0.00035125163,0.98418343,0.002412305,0.00006284809,0.00018699453,0.0009378429,0.0013626287,0.006542374],"genre_scores_gemma":[0.09081318,0.0013105056,0.89997023,0.0007428769,0.000114103714,0.00047003347,0.004130709,0.00037646058,0.002071911],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.992455,0.0025661173,0.0013691639,0.0007652216,0.0025711309,0.00027331515],"domain_scores_gemma":[0.9823815,0.010789204,0.0013615219,0.0030259707,0.002156874,0.0002849066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008483205,0.0010075604,0.00070033554,0.0048334403,0.001050676,0.0055929953,0.0024808436,0.0019034125,0.003347977],"category_scores_gemma":[0.020474894,0.0009013695,0.0022746301,0.00427046,0.0022906163,0.008270483,0.0049805134,0.0025668826,0.0009112114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004854722,0.00012598245,0.002176792,0.00083193555,0.00017100327,0.0006273909,0.0017059093,0.066429496,0.0032043979,0.8414858,0.00578244,0.07741033],"study_design_scores_gemma":[0.000029200532,0.00003272298,0.00060587336,0.00048844476,0.00011803039,0.0003335717,0.00040337665,0.17641833,0.005425286,0.6982189,0.11786517,0.00006120114],"about_ca_topic_score_codex":0.006519077,"about_ca_topic_score_gemma":0.0049007735,"teacher_disagreement_score":0.008483205,"about_ca_system_score_codex":0.0033327658,"about_ca_system_score_gemma":0.0047962004,"threshold_uncertainty_score":0.044864},"labels":[],"label_agreement":null},{"id":"W2136961946","doi":"10.1186/1471-2105-11-441","title":"Relations as patterns: bridging the gap between OBO and OWL","year":2010,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":66,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Ontology; Open Biomedical Ontologies; Web Ontology Language; Ontology components; Semantics (computer science); Ontology language; OWL-S; Description logic; Programming language; Software; Automated reasoning; Protégé; Natural language processing; Bridging (networking); Information retrieval; Semantic Web; Artificial intelligence; Semantic Web Stack","score_opus":0.02210392464390946,"score_gpt":0.2746213768207904,"score_spread":0.25251745217688093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136961946","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004686298,0.00046237215,0.9837969,0.0019654909,0.000122056124,0.00020625089,0.00032352188,0.0022202583,0.006216838],"genre_scores_gemma":[0.061391696,0.0011925237,0.92912185,0.0014386587,0.00013177654,0.00044766563,0.0013144318,0.0018160265,0.0031453401],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98904896,0.003561922,0.0017166854,0.0012008209,0.0037854302,0.0006861262],"domain_scores_gemma":[0.98257154,0.007905746,0.0015079526,0.005675954,0.0017297874,0.00060898124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0109055415,0.000959811,0.001026063,0.002868053,0.0013848941,0.0060486803,0.0030519078,0.0018029219,0.0036620551],"category_scores_gemma":[0.022963013,0.001523819,0.00215761,0.0038683012,0.005359548,0.014903774,0.0070774877,0.0044196937,0.0015091562],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001878035,0.00013515449,0.0027557795,0.00096451805,0.00009246308,0.0005842083,0.004007187,0.0029410685,0.006007082,0.7302299,0.012202287,0.2398926],"study_design_scores_gemma":[0.00007425381,0.000113838854,0.0016899769,0.0010342924,0.000118699936,0.0012425139,0.0014212951,0.0400511,0.008250876,0.47270635,0.47317,0.00012675539],"about_ca_topic_score_codex":0.008431224,"about_ca_topic_score_gemma":0.00782435,"teacher_disagreement_score":0.0109055415,"about_ca_system_score_codex":0.0020888448,"about_ca_system_score_gemma":0.0042534177,"threshold_uncertainty_score":0.057674766},"labels":[],"label_agreement":null},{"id":"W2136992683","doi":"10.1186/1471-2105-7-41","title":"Discovering semantic features in the literature: a foundation for building functional associations","year":2006,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":92,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Cluster analysis; Computer science; Context (archaeology); DNA microarray; Non-negative matrix factorization; Data mining; Semantic similarity; Interpretation (philosophy); Information retrieval; Computational biology; Natural language processing; Gene; Artificial intelligence; Biology; Matrix decomposition; Gene expression; Genetics","score_opus":0.016823192533370328,"score_gpt":0.2701781252555171,"score_spread":0.25335493272214676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2136992683","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13231748,0.0053078183,0.8036925,0.006266258,0.00041899158,0.0016830812,0.026575888,0.004960177,0.01877782],"genre_scores_gemma":[0.26701462,0.001964879,0.7094691,0.00035235204,0.00035849985,0.0010506507,0.018530136,0.00025062,0.0010092201],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9976655,0.0004730191,0.00044709272,0.0006937452,0.0005813321,0.00013928942],"domain_scores_gemma":[0.98825175,0.006090936,0.0014745492,0.0012407982,0.0024120337,0.0005300347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030305667,0.0010934476,0.0008559096,0.03348627,0.0019977677,0.0039240182,0.0013904487,0.0013879593,0.0041987174],"category_scores_gemma":[0.018524738,0.00052943005,0.002111416,0.016204134,0.0015858707,0.008341298,0.0022459957,0.0011638979,0.0020170866],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052289997,0.00078331836,0.05005751,0.0032236017,0.0005296736,0.0017811096,0.004183557,0.007317566,0.015157989,0.12435585,0.018518703,0.7735682],"study_design_scores_gemma":[0.00016905193,0.00044231102,0.04330701,0.0028466368,0.0009504606,0.0035724784,0.00604994,0.14399542,0.010947515,0.6435935,0.1438259,0.00029977504],"about_ca_topic_score_codex":0.00260639,"about_ca_topic_score_gemma":0.0035172678,"teacher_disagreement_score":0.03348627,"about_ca_system_score_codex":0.0013264966,"about_ca_system_score_gemma":0.0029854379,"threshold_uncertainty_score":0.016027391},"labels":[],"label_agreement":null},{"id":"W2137313573","doi":"10.1109/titb.2005.847188","title":"A Knowledge Creation Info-Structure to Acquire and Crystallize the Tacit Knowledge of Health-Care Experts","year":2005,"lang":"en","type":"article","venue":"IEEE Transactions on Information Technology in Biomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":80,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Izaak Walton Killam Health Centre; Dalhousie University","funders":"","keywords":"Tacit knowledge; Knowledge management; Experiential knowledge; Health care; Explicit knowledge; Knowledge value chain; Reuse; Computer science; Organizational learning; Engineering; Political science; Epistemology","score_opus":0.008062049210254349,"score_gpt":0.28837475035248994,"score_spread":0.2803127011422356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2137313573","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022315782,0.00018673552,0.9478996,0.0036888542,0.000040427545,0.00026314097,0.0006412876,0.0016862452,0.02327798],"genre_scores_gemma":[0.17284083,0.00026788228,0.82135814,0.000310168,0.00002607343,0.000218111,0.0013137235,0.0000811665,0.0035839283],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99891186,0.00041293245,0.00012886383,0.0002514144,0.00023281186,0.00006211114],"domain_scores_gemma":[0.9959818,0.0018725509,0.00033447967,0.0012121581,0.00041652264,0.00018256555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025468257,0.00042597819,0.00033663685,0.002495664,0.001363824,0.004454326,0.0013435321,0.0012314363,0.005781466],"category_scores_gemma":[0.0062821484,0.00045394388,0.0010893003,0.0019715063,0.0021104298,0.007013203,0.0032342563,0.0014628934,0.0012430655],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014894865,0.00029369403,0.0061057233,0.0004865679,0.00010977443,0.00069911324,0.006033928,0.020111775,0.0065159644,0.69488555,0.008010937,0.25659803],"study_design_scores_gemma":[0.00012647177,0.00020554822,0.006790956,0.0008137498,0.00023589526,0.0017620877,0.0027921346,0.2533755,0.027123883,0.51097864,0.19566578,0.00012940455],"about_ca_topic_score_codex":0.0032846779,"about_ca_topic_score_gemma":0.0040569785,"teacher_disagreement_score":0.005781466,"about_ca_system_score_codex":0.0012705959,"about_ca_system_score_gemma":0.0032185551,"threshold_uncertainty_score":0.019340932},"labels":[],"label_agreement":null},{"id":"W2137378906","doi":"10.1186/2041-1480-2-s2-s3","title":"HyQue: evaluating hypotheses using Semantic Web technologies","year":2011,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"National Human Genome Research Institute; Natural Sciences and Engineering Research Council of Canada; National Institutes of Health; National Science Foundation","keywords":"Computer science; SPARQL; RDF; Semantic Web; Information retrieval; Inference; Linked data; Data science; Artificial intelligence","score_opus":0.09708492824450848,"score_gpt":0.3320480927395311,"score_spread":0.23496316449502264,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2137378906","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021088067,0.0015801674,0.9054988,0.0037994995,0.00024489718,0.0020499614,0.012629357,0.041008048,0.012101292],"genre_scores_gemma":[0.11512449,0.0010291408,0.86074543,0.0011255238,0.00016558073,0.0010006048,0.017508728,0.0016778206,0.0016226232],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98261607,0.007123234,0.0021318316,0.002028815,0.0057316297,0.00036850752],"domain_scores_gemma":[0.9346653,0.050792653,0.0026581644,0.005146859,0.0057337554,0.0010032686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028175527,0.0019824451,0.0013675673,0.010490992,0.001813836,0.008962177,0.003944346,0.003058686,0.011406697],"category_scores_gemma":[0.06339495,0.0008508782,0.0036268942,0.0036566711,0.0035216245,0.012502115,0.0053802994,0.0018341595,0.002442693],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015638162,0.001387811,0.018698223,0.0061943484,0.0012755397,0.0024746917,0.0033031106,0.07367686,0.014144234,0.26168752,0.07171556,0.5438783],"study_design_scores_gemma":[0.00038274075,0.00030052112,0.0040282435,0.0010704246,0.00036676627,0.00061466795,0.0015390606,0.3361509,0.024682267,0.5215404,0.10908756,0.0002365122],"about_ca_topic_score_codex":0.004854417,"about_ca_topic_score_gemma":0.004340675,"teacher_disagreement_score":0.028175527,"about_ca_system_score_codex":0.002385663,"about_ca_system_score_gemma":0.0042874445,"threshold_uncertainty_score":0.14900821},"labels":[],"label_agreement":null},{"id":"W2138088864","doi":"10.1093/database/bau060","title":"A controlled vocabulary for pathway entities and events","year":2014,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Human Genome Research Institute; National Institutes of Health; European Commission; European Bioinformatics Institute","keywords":"Computer science; Readability; Consistency (knowledge bases); Vocabulary; Information retrieval; Controlled vocabulary; World Wide Web; Protocol (science); Database; State (computer science); Programming language; Artificial intelligence; Linguistics","score_opus":0.01190032490260914,"score_gpt":0.25635823974965316,"score_spread":0.24445791484704402,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138088864","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035642518,0.00089261605,0.9192972,0.001213357,0.00093051675,0.0022479093,0.044759348,0.010687072,0.016407719],"genre_scores_gemma":[0.04835242,0.0018428749,0.8267744,0.0015212414,0.0005966555,0.005417611,0.10342747,0.0035066402,0.008560718],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9905716,0.001867594,0.0036992014,0.0019086319,0.0016020016,0.00035102968],"domain_scores_gemma":[0.9857248,0.0052800393,0.001574148,0.0036926498,0.0030351295,0.00069329125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009460259,0.0017432766,0.0018478868,0.008702577,0.0030456348,0.008516918,0.0036277613,0.0027424595,0.011223582],"category_scores_gemma":[0.016498912,0.0014234667,0.0032362232,0.0071583237,0.0033567187,0.012284934,0.004796132,0.00509222,0.007988214],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002546641,0.00011551418,0.001006044,0.0016207618,0.00007884914,0.0006099498,0.0020467394,0.0035563994,0.010328185,0.8717358,0.05236252,0.056284618],"study_design_scores_gemma":[0.0000942851,0.000068568166,0.0004078529,0.00045155192,0.00008098776,0.0006512568,0.0004011137,0.0070856097,0.005009214,0.14131704,0.844343,0.00008955265],"about_ca_topic_score_codex":0.008645179,"about_ca_topic_score_gemma":0.007129062,"teacher_disagreement_score":0.011223582,"about_ca_system_score_codex":0.0034834608,"about_ca_system_score_gemma":0.0099015925,"threshold_uncertainty_score":0.050031185},"labels":[],"label_agreement":null},{"id":"W2138586647","doi":"10.1145/2166896.2166910","title":"Using semantic web technology to support ICD-11 textual definitions authoring","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institute for Theoretical Astrophysics","keywords":"Computer science; SNOMED CT; Linked data; RDF; Semantic Web; Unified Medical Language System; Crowdsourcing; Information retrieval; World Wide Web; Annotation; Terminology; Artificial intelligence","score_opus":0.11127515826057258,"score_gpt":0.32784755234657514,"score_spread":0.21657239408600257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2138586647","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012097121,0.00015542559,0.97271967,0.00063499634,0.00009004458,0.0006403769,0.00056779064,0.007936114,0.005158387],"genre_scores_gemma":[0.075657494,0.00024894506,0.9176609,0.0002126424,0.00006372838,0.00044841578,0.0023108786,0.0010203526,0.0023767187],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99192715,0.0041269655,0.0011593921,0.0008356508,0.0017450592,0.0002056942],"domain_scores_gemma":[0.9688885,0.020291742,0.0015798886,0.0049471017,0.003673234,0.00061943126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014869599,0.0008592645,0.0007809753,0.0030235455,0.0010131641,0.0041139363,0.0018731386,0.0015081061,0.0037676434],"category_scores_gemma":[0.027746301,0.00064940914,0.0012436173,0.0020321608,0.0013595426,0.005794689,0.0036685034,0.0019589625,0.0016689174],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009956552,0.0013342602,0.008147361,0.0020825462,0.00040332862,0.0035048719,0.011786771,0.038242582,0.056337573,0.13241489,0.03218207,0.71256816],"study_design_scores_gemma":[0.0005065729,0.00038844513,0.0024787642,0.00082156155,0.0003062258,0.0021647047,0.0023471299,0.36700642,0.12105633,0.19076355,0.31183037,0.0003298672],"about_ca_topic_score_codex":0.0008811833,"about_ca_topic_score_gemma":0.0015533105,"teacher_disagreement_score":0.014869599,"about_ca_system_score_codex":0.00092209055,"about_ca_system_score_gemma":0.0017279119,"threshold_uncertainty_score":0.07863891},"labels":[],"label_agreement":null},{"id":"W2139259976","doi":"10.1186/1471-2105-4-11","title":"PreBIND and Textomy – mining the biomedical literature for protein-protein interactions using a support vector machine","year":2003,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":340,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; National Research Council Canada; Lunenfeld-Tanenbaum Research Institute","funders":"Canadian Institutes of Health Research; Directorate for Biological Sciences; Genome Canada","keywords":"Computer science; Support vector machine; Set (abstract data type); Task (project management); Interaction information; Precision and recall; Recall; Data curation; Interaction network; Artificial intelligence; Protein–protein interaction; Data mining; Machine learning; Database; Information retrieval; Biology","score_opus":0.027212083889263005,"score_gpt":0.29671629648987846,"score_spread":0.26950421260061547,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139259976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3048123,0.03055897,0.44912943,0.0064884583,0.0010801342,0.004479551,0.13948278,0.04922898,0.014739333],"genre_scores_gemma":[0.2728015,0.005532228,0.60241234,0.0007535659,0.00039767078,0.002010341,0.10734342,0.0005317177,0.00821724],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99782807,0.0004170439,0.0005289041,0.00053281244,0.0005994194,0.00009379743],"domain_scores_gemma":[0.9913651,0.004296163,0.0012208937,0.0006881528,0.0021459528,0.00028381593],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036640374,0.0011159106,0.0011307918,0.016906857,0.0006997165,0.0021747882,0.0013259301,0.000947494,0.008280493],"category_scores_gemma":[0.012391887,0.0003690659,0.0013028565,0.0070893094,0.00054024585,0.0021746496,0.0014353298,0.0008191144,0.0046075247],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094555103,0.00039896808,0.02033173,0.0066350196,0.0004595455,0.0010750735,0.0006788184,0.005066533,0.037361152,0.002692855,0.052853905,0.87150097],"study_design_scores_gemma":[0.000700597,0.002955172,0.092262544,0.0023017458,0.0013203892,0.0075510843,0.002627303,0.31728104,0.17062686,0.033459272,0.36847067,0.00044332398],"about_ca_topic_score_codex":0.0022097067,"about_ca_topic_score_gemma":0.0037471433,"teacher_disagreement_score":0.016906857,"about_ca_system_score_codex":0.00080825883,"about_ca_system_score_gemma":0.003111933,"threshold_uncertainty_score":0.02770096},"labels":[],"label_agreement":null},{"id":"W2139609609","doi":"10.3389/neuro.11.029.2009","title":"Automated recognition of brain region mentions in neuroscience literature","year":2009,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of General Medical Sciences; Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; National Institutes of Health; Michael Smith Health Research BC","keywords":"Computer science; Natural language processing; Artificial intelligence; Context (archaeology); Lemmatisation; Set (abstract data type); Feature (linguistics); Vocabulary; Conditional random field; Matching (statistics); Precision and recall; Pattern recognition (psychology); Linguistics; Biology","score_opus":0.014990091980470069,"score_gpt":0.2631887894516649,"score_spread":0.24819869747119483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2139609609","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40782145,0.03288715,0.29450157,0.0029735127,0.0012815192,0.002175763,0.19868904,0.03273976,0.026930206],"genre_scores_gemma":[0.357643,0.0065130927,0.47630367,0.00025025904,0.000503402,0.0008829034,0.15208776,0.0008933388,0.004922552],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967534,0.0005525972,0.0009598366,0.0009240637,0.0006464608,0.00016350702],"domain_scores_gemma":[0.96366715,0.021435883,0.0057143834,0.002907351,0.0056808926,0.0005943435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004713159,0.0008591122,0.0011334594,0.04014968,0.0013506416,0.0022037064,0.0013302864,0.0010543424,0.0071771266],"category_scores_gemma":[0.027007623,0.0007391298,0.0011448588,0.021304645,0.00079797755,0.0044155247,0.0026720339,0.0007180181,0.0048294626],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074059196,0.00017020982,0.049643826,0.014455472,0.00027720392,0.0039262036,0.004636947,0.004286392,0.09146213,0.01234389,0.076914236,0.7411429],"study_design_scores_gemma":[0.00025242535,0.0005689421,0.2098987,0.0033124364,0.001116171,0.015418832,0.005362512,0.07422978,0.0936241,0.030436678,0.565322,0.0004574762],"about_ca_topic_score_codex":0.003949644,"about_ca_topic_score_gemma":0.0071728094,"teacher_disagreement_score":0.04014968,"about_ca_system_score_codex":0.0011259527,"about_ca_system_score_gemma":0.0025463703,"threshold_uncertainty_score":0.024925888},"labels":[],"label_agreement":null},{"id":"W2140001356","doi":"10.1093/bioinformatics/btp259","title":"Application and evaluation of automated semantic annotation of gene expression experiments","year":2009,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of General Medical Sciences; Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; National Institutes of Health; Michael Smith Health Research BC","keywords":"Computer science; Unified Medical Language System; Annotation; Natural language processing; Documentation; Information retrieval; Controlled vocabulary; Software; Expression (computer science); Pipeline (software); Context (archaeology); Artificial intelligence; Source code; Programming language","score_opus":0.023997876135531407,"score_gpt":0.326307267117323,"score_spread":0.3023093909817916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140001356","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46944362,0.0040061567,0.36867893,0.002153545,0.0009998892,0.0025848874,0.02159773,0.11808798,0.012447293],"genre_scores_gemma":[0.40712982,0.0005999284,0.532751,0.00074894744,0.00013561497,0.0012580232,0.05063585,0.0035414076,0.0031994076],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9731358,0.013207208,0.0020542552,0.005260631,0.005697173,0.0006448645],"domain_scores_gemma":[0.9398437,0.035362393,0.0020940257,0.009431545,0.0123783285,0.00089008396],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029345382,0.0022710834,0.0013606638,0.005925643,0.001903368,0.00335349,0.0033679425,0.0029623345,0.0042124866],"category_scores_gemma":[0.07339972,0.0005874386,0.0019505899,0.0032050977,0.0017656646,0.0037620165,0.004196067,0.0015755616,0.0024192077],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007853603,0.0028359804,0.040833563,0.009749242,0.0019020897,0.0015108454,0.0052680513,0.06268299,0.15610409,0.012067199,0.06640964,0.63278264],"study_design_scores_gemma":[0.00092501414,0.0021414254,0.049518757,0.0008201449,0.0010867455,0.0013500873,0.0023681961,0.61190253,0.2207765,0.016221464,0.0924864,0.00040277367],"about_ca_topic_score_codex":0.007700226,"about_ca_topic_score_gemma":0.0063476106,"teacher_disagreement_score":0.029345382,"about_ca_system_score_codex":0.0031486598,"about_ca_system_score_gemma":0.004118334,"threshold_uncertainty_score":0.15519506},"labels":[],"label_agreement":null},{"id":"W2140080033","doi":"10.1186/2041-1480-2-s2-s6","title":"Scalable representations of diseases in biomedical ontologies","year":2011,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Francis Xavier University; University of Toronto","funders":"Deutsche Forschungsgemeinschaft","keywords":"Ontology; Computer science; SNOMED CT; Open Biomedical Ontologies; Process ontology; Domain (mathematical analysis); Process (computing); Hierarchy; Soundness; Ontology components; Information retrieval; Upper ontology; Natural language processing; Data science; Artificial intelligence; Ontology alignment; Conceptualization; Terminology; Epistemology; Programming language; Linguistics; Mathematics","score_opus":0.029547549609980482,"score_gpt":0.29836785880152145,"score_spread":0.268820309191541,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2140080033","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017037619,0.001059196,0.9589719,0.002090826,0.00011427475,0.00047534,0.0067945756,0.008516587,0.0049395324],"genre_scores_gemma":[0.17591701,0.001477583,0.8038382,0.0003451351,0.00009344465,0.00048243807,0.015144209,0.00039682945,0.0023051826],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9980665,0.0005092417,0.00036692494,0.00030379603,0.00065753126,0.000096121694],"domain_scores_gemma":[0.99534565,0.001998694,0.0004553294,0.0015067873,0.00053024344,0.0001633116],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038057975,0.00065028796,0.00093556097,0.0037413517,0.0010180168,0.0055250847,0.0021258167,0.0013820344,0.004143581],"category_scores_gemma":[0.012832739,0.0005019416,0.0022859427,0.0054218005,0.00093866716,0.0065045436,0.004783428,0.0014733368,0.0011847499],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034183377,0.00029721728,0.007218545,0.0022676014,0.00043139327,0.001416351,0.0022708983,0.18888946,0.008606938,0.29431736,0.034848902,0.45909345],"study_design_scores_gemma":[0.000098007666,0.00005833504,0.0018465373,0.00040673773,0.00025814783,0.00059473945,0.00062290573,0.51398385,0.006013542,0.40227354,0.07378655,0.000057150468],"about_ca_topic_score_codex":0.006551978,"about_ca_topic_score_gemma":0.011396082,"teacher_disagreement_score":0.006551978,"about_ca_system_score_codex":0.0024321158,"about_ca_system_score_gemma":0.0028855407,"threshold_uncertainty_score":0.020127237},"labels":[],"label_agreement":null},{"id":"W2141297584","doi":"10.1038/ng.1054","title":"Toward interoperable bioscience data","year":2012,"lang":"en","type":"article","venue":"Nature Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":468,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Cancer Institute; National Heart, Lung, and Blood Institute; National Human Genome Research Institute; National Center for Research Resources; National Institute of General Medical Sciences; National Institute of Mental Health; Biotechnology and Biological Sciences Research Council","keywords":"Interoperability; Biology; Data science; Knowledge management; Data curation; Research data; World Wide Web; Computer science","score_opus":0.050154971685411776,"score_gpt":0.3283915047080819,"score_spread":0.27823653302267015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141297584","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003754554,0.0010709994,0.9586614,0.019514102,0.0006145621,0.00040093673,0.0025777372,0.0038742174,0.009531468],"genre_scores_gemma":[0.0371897,0.0017968249,0.9360473,0.005155799,0.00049465563,0.0007945852,0.014688757,0.0008981344,0.0029342796],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9410763,0.024195552,0.009593097,0.008508712,0.014379534,0.0022469233],"domain_scores_gemma":[0.8088862,0.03817484,0.006903705,0.1155645,0.022435376,0.008035413],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.09935374,0.0015589821,0.0027770465,0.011836438,0.0040406077,0.019215234,0.010122755,0.0063484046,0.0038031172],"category_scores_gemma":[0.113306284,0.0020237947,0.004433376,0.0147191845,0.008132686,0.0443039,0.038855687,0.0149475075,0.0029401807],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013101556,0.0002053959,0.0027866564,0.0006373949,0.00022302427,0.00032877622,0.002584997,0.004512328,0.00258558,0.85909677,0.022886135,0.10402195],"study_design_scores_gemma":[0.000046871584,0.000042319436,0.0006206969,0.0006108481,0.00011818009,0.00017963431,0.0010777469,0.012577693,0.0028710577,0.76883304,0.21294743,0.000074460295],"about_ca_topic_score_codex":0.0054552034,"about_ca_topic_score_gemma":0.0047261594,"teacher_disagreement_score":0.9898772,"about_ca_system_score_codex":0.004932616,"about_ca_system_score_gemma":0.018728321,"threshold_uncertainty_score":0.525439},"labels":[],"label_agreement":null},{"id":"W2141363231","doi":"10.1093/bioinformatics/btq416","title":"OWL2Perl: creating Perl modules from OWL class definitions","year":2010,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital","funders":"Natural Sciences and Engineering Research Council of Canada; Genome Canada; Canarie","keywords":"Perl; Computer science; RDF; Class (philosophy); Programming language; World Wide Web; Software; Resource (disambiguation); License; Semantic Web; Software engineering; Artificial intelligence; Operating system","score_opus":0.024532266281811578,"score_gpt":0.25449899228227674,"score_spread":0.22996672600046517,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2141363231","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024334297,0.00012890642,0.8960788,0.00045946,0.00020709414,0.00054511824,0.0087794745,0.084096,0.0072717196],"genre_scores_gemma":[0.040303774,0.00045671477,0.86772597,0.0009061791,0.00014035638,0.001363511,0.04687598,0.03197961,0.010247943],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964833,0.0007774074,0.00045847445,0.00054959615,0.0014810426,0.00024999934],"domain_scores_gemma":[0.9934534,0.002669481,0.00067313394,0.0016564472,0.0012715074,0.00027615431],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007455349,0.0018020999,0.0008197199,0.0028038346,0.000989591,0.0031686875,0.004368017,0.0013343387,0.021556273],"category_scores_gemma":[0.0105353175,0.0018147686,0.0022373123,0.0018778811,0.0011526314,0.0064134495,0.0034764446,0.0030502034,0.012664567],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005823979,0.0005662879,0.0032612616,0.0030368424,0.00041783892,0.001390149,0.0015976684,0.012751428,0.02393344,0.13202535,0.33270976,0.48772752],"study_design_scores_gemma":[0.00018510401,0.00008820302,0.0010753631,0.0005656466,0.00013257559,0.0010350225,0.00024300344,0.12750332,0.056428958,0.068613335,0.74392366,0.0002057707],"about_ca_topic_score_codex":0.0029678985,"about_ca_topic_score_gemma":0.003619972,"teacher_disagreement_score":0.021556273,"about_ca_system_score_codex":0.001130185,"about_ca_system_score_gemma":0.0024371315,"threshold_uncertainty_score":0.07211298},"labels":[],"label_agreement":null},{"id":"W2142407957","doi":"10.1109/tkde.2010.152","title":"A Machine Learning Approach for Identifying Disease-Treatment Relations in Short Texts","year":2010,"lang":"en","type":"article","venue":"IEEE Transactions on Knowledge and Data Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Domain (mathematical analysis); Field (mathematics); Set (abstract data type); Health care; Machine learning; Artificial intelligence; Dissemination; Information extraction; Data science","score_opus":0.036150843997362125,"score_gpt":0.3036804357149973,"score_spread":0.2675295917176352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2142407957","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029543152,0.0050776666,0.9463759,0.003036985,0.0005575659,0.0009008871,0.007989472,0.0026561108,0.0038621905],"genre_scores_gemma":[0.20764992,0.00177557,0.77022076,0.00085681415,0.0012234941,0.0015860393,0.013496944,0.00012215662,0.0030682816],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99557686,0.0016490916,0.0007966117,0.001066242,0.00078206963,0.00012914646],"domain_scores_gemma":[0.98871243,0.009198558,0.0008685,0.00041626894,0.00067013566,0.00013406368],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037036338,0.0014498367,0.0010760013,0.008528781,0.0012503392,0.0017027095,0.0012663895,0.002082722,0.0032395334],"category_scores_gemma":[0.014033899,0.00041736054,0.0015403581,0.006561076,0.0008854509,0.0032748578,0.001145486,0.001897935,0.0021485945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005914831,0.0007850201,0.010086382,0.0016916329,0.00049745396,0.0008792362,0.0009176021,0.021205531,0.015361419,0.01547331,0.015266018,0.91724485],"study_design_scores_gemma":[0.00025400944,0.00084955094,0.01765995,0.0005058808,0.000660889,0.0021348891,0.000790077,0.809005,0.015308147,0.094279625,0.05831947,0.0002324901],"about_ca_topic_score_codex":0.0024617983,"about_ca_topic_score_gemma":0.003383113,"teacher_disagreement_score":0.008528781,"about_ca_system_score_codex":0.0010037238,"about_ca_system_score_gemma":0.0016850717,"threshold_uncertainty_score":0.01958692},"labels":[],"label_agreement":null},{"id":"W2143572567","doi":"10.1109/cbms.2005.74","title":"Medical Knowledge Morphing: Towards Case-Specific Integration of Heterogeneous Medical Knowledge Resources","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Modalities; Knowledge management; Medical knowledge; Computer science; Tacit knowledge; Experiential knowledge; Procedural knowledge; Knowledge engineering; Personal knowledge management; Task (project management); Body of knowledge; Domain knowledge; Knowledge-based systems; Knowledge integration; Explicit knowledge; Resource (disambiguation); Knowledge value chain; Organizational learning; Medicine; Engineering","score_opus":0.03156646501583091,"score_gpt":0.31605956508973654,"score_spread":0.2844931000739056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143572567","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013742652,0.00039899867,0.978588,0.0016504176,0.000036169684,0.00030545003,0.00047702703,0.0017548042,0.003046467],"genre_scores_gemma":[0.1393725,0.0005214601,0.8570725,0.00039898767,0.000034525834,0.00018342645,0.0012900231,0.00018389395,0.00094274856],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958352,0.0017396895,0.0005030105,0.0007601732,0.0010112504,0.0001507823],"domain_scores_gemma":[0.9957463,0.0022798681,0.00044938223,0.0010742557,0.00029589853,0.00015418416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067472868,0.0008909909,0.0009549459,0.004729314,0.0008296544,0.004648943,0.002164053,0.0015925849,0.0023539786],"category_scores_gemma":[0.015612649,0.000747271,0.0017557248,0.0040066247,0.0019411693,0.00891424,0.0057343584,0.0014924507,0.0005826295],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034348204,0.00035769763,0.009994507,0.0010096126,0.0003945424,0.003893066,0.004301655,0.04205634,0.022863826,0.206459,0.011692934,0.6966334],"study_design_scores_gemma":[0.000086706554,0.00016667969,0.0030637635,0.0004952486,0.00047434468,0.0052459375,0.00232485,0.4384686,0.037500035,0.41193277,0.10006788,0.00017310934],"about_ca_topic_score_codex":0.002197853,"about_ca_topic_score_gemma":0.002896836,"teacher_disagreement_score":0.0067472868,"about_ca_system_score_codex":0.0009984067,"about_ca_system_score_gemma":0.0017042819,"threshold_uncertainty_score":0.035683513},"labels":[],"label_agreement":null},{"id":"W2144427809","doi":"10.1186/2041-1480-5-14","title":"The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2014,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":272,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"Instituto de Salud Carlos III; Natural Sciences and Engineering Research Council of Canada; National Science Foundation; University of Texas at El Paso; Canarie; European Federation of Pharmaceutical Industries and Associations; National Aeronautics and Space Administration","keywords":"Computer science; Ontology; License; Simple (philosophy); Semantic Web; Data science; World Wide Web; Open Biomedical Ontologies; Information retrieval; Ontology-based data integration; Suggested Upper Merged Ontology","score_opus":0.04075967433146419,"score_gpt":0.36765102721213694,"score_spread":0.3268913528806727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144427809","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029437284,0.0026257616,0.9216031,0.0069069257,0.0011065607,0.0007151,0.006366709,0.004321621,0.053410478],"genre_scores_gemma":[0.048317276,0.005056381,0.9077468,0.003244189,0.00085996394,0.0014109212,0.019583639,0.0010405992,0.012740129],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99555004,0.0011512298,0.00086071965,0.00056173437,0.0015809466,0.00029537224],"domain_scores_gemma":[0.994873,0.0015728689,0.0005043746,0.0014331832,0.0010888935,0.00052764785],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006278325,0.00077744824,0.0009888267,0.0064556003,0.002473816,0.006231953,0.0020781257,0.0023223679,0.0050084637],"category_scores_gemma":[0.0087327445,0.0007536242,0.0023105587,0.0072078635,0.0031578136,0.01041476,0.006162736,0.0034558417,0.0035007761],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000033467677,0.00006221025,0.0005294012,0.0005039067,0.000064871296,0.00020205663,0.00064028177,0.0012006862,0.0019454968,0.88580835,0.03858476,0.070424534],"study_design_scores_gemma":[0.00002712851,0.000022362341,0.00053895917,0.00049055077,0.000054766402,0.00050972763,0.0003101719,0.0073356694,0.0014297451,0.37626243,0.6129711,0.000047263173],"about_ca_topic_score_codex":0.0053643268,"about_ca_topic_score_gemma":0.0058559286,"teacher_disagreement_score":0.0064556003,"about_ca_system_score_codex":0.003677141,"about_ca_system_score_gemma":0.012328218,"threshold_uncertainty_score":0.033203304},"labels":[],"label_agreement":null},{"id":"W2144631782","doi":"10.1093/bib/bbn056","title":"Towards pharmacogenomics knowledge discovery with the semantic web","year":2009,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":64,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Pharmacogenomics; Computer science; Semantic Web; Semantics (computer science); Knowledge extraction; Data science; Open Biomedical Ontologies; XML; World Wide Web; Social Semantic Web; Bioinformatics; Artificial intelligence; Biology; OWL-S","score_opus":0.011075004144815638,"score_gpt":0.2613136465209089,"score_spread":0.2502386423760932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2144631782","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041597635,0.0058327466,0.9700724,0.010340882,0.00025776096,0.0001402513,0.0006181616,0.0012175493,0.00736057],"genre_scores_gemma":[0.042781267,0.010778562,0.9385916,0.0025265736,0.00034370692,0.00023015423,0.0024295968,0.00015292151,0.002165565],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935294,0.0029651716,0.0006756321,0.0006044985,0.0019651055,0.00026008222],"domain_scores_gemma":[0.99288905,0.0045735845,0.00032940615,0.0009894137,0.0010224194,0.00019615957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01283111,0.0010486037,0.0017876143,0.006551388,0.001428071,0.008841988,0.0026446572,0.0036841983,0.0023174568],"category_scores_gemma":[0.013412484,0.0011055282,0.0033277741,0.006491261,0.004169179,0.016778227,0.005124772,0.005126009,0.0014331527],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001695011,0.00025499985,0.0016687093,0.001509245,0.00044717218,0.0011275483,0.0014160329,0.021106021,0.0042134547,0.6719092,0.018458454,0.27771956],"study_design_scores_gemma":[0.0000368506,0.000018336103,0.0002818403,0.0003448876,0.00010073697,0.00031199897,0.00045495157,0.0544281,0.0023023502,0.8522679,0.08940114,0.000050920506],"about_ca_topic_score_codex":0.003850713,"about_ca_topic_score_gemma":0.004497609,"teacher_disagreement_score":0.01283111,"about_ca_system_score_codex":0.0019050445,"about_ca_system_score_gemma":0.004017047,"threshold_uncertainty_score":0.06785822},"labels":[],"label_agreement":null},{"id":"W2145409264","doi":"10.1016/j.artmed.2010.04.012","title":"A four stage approach for ontology-based health information system design","year":2010,"lang":"en","type":"article","venue":"Artificial Intelligence in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":65,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Wilfrid Laurier University; University of Ottawa","funders":"","keywords":"Computer science; Ontology; Ontology-based data integration; Process ontology; Usability; Suggested Upper Merged Ontology; Upper ontology; Ontology alignment; Software engineering; Information retrieval; Domain knowledge; Human–computer interaction","score_opus":0.10285494844143475,"score_gpt":0.35180835181773656,"score_spread":0.2489534033763018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145409264","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015424415,0.00005008496,0.99453133,0.0006327481,0.00001812323,0.000514015,0.00009360021,0.0004129843,0.0022046422],"genre_scores_gemma":[0.018199295,0.00006929363,0.9793802,0.00014096579,0.000007753428,0.0003218315,0.0002687603,0.00006774444,0.0015441283],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9939812,0.0019109204,0.00076987105,0.0005672146,0.0024372265,0.00033358825],"domain_scores_gemma":[0.993816,0.0023051896,0.000282623,0.0008858586,0.0023504929,0.00035979057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008892523,0.0009895944,0.0010862936,0.0034077144,0.0024932222,0.0063577057,0.0035970805,0.0023133073,0.0069401176],"category_scores_gemma":[0.010331838,0.0016290381,0.0032257943,0.0021938183,0.0027351452,0.006088495,0.004485978,0.0029546563,0.0020639836],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039863097,0.0008296282,0.0040075183,0.0015691415,0.00074801926,0.0011706164,0.006926854,0.04322686,0.03943657,0.51624954,0.015411025,0.37002558],"study_design_scores_gemma":[0.0002112682,0.00046903943,0.0018061654,0.0005683786,0.000782975,0.0010683608,0.0021177242,0.4206087,0.040839612,0.41356504,0.11767846,0.00028425315],"about_ca_topic_score_codex":0.009300473,"about_ca_topic_score_gemma":0.012798191,"teacher_disagreement_score":0.009300473,"about_ca_system_score_codex":0.0027381293,"about_ca_system_score_gemma":0.00799896,"threshold_uncertainty_score":0.04702872},"labels":[],"label_agreement":null},{"id":"W2145428772","doi":"10.1186/1471-2105-12-481","title":"U-Compare bio-event meta-service: compatible BioNLP event extraction services","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Brain Injury Research Center; U.S. National Library of Medicine; National Institute of General Medical Sciences; Precursory Research for Embryonic Science and Technology; National Institute on Drug Abuse; Directorate for Biological Sciences; National Institutes of Health; Fonds Wetenschappelijk Onderzoek; China Scholarship Council; Joint Information Systems Committee; Vlaamse regering; Biotechnology and Biological Sciences Research Council; University of Tokyo; Korea Institute of Science and Technology Information; University of Massachusetts Amherst; Korea Institute of Science and Technology; Norges Teknisk-Naturvitenskapelige Universitet; Japan Society for the Promotion of Science; Academy of Finland; Ministry of Education, Culture, Sports, Science and Technology; National Institute of Advanced Industrial Science and Technology","keywords":"Computer science; Event (particle physics); Biomedical text mining; Interoperability; Service (business); Task (project management); Data mining; Extraction (chemistry); Information extraction; Information retrieval; Text mining; World Wide Web; Systems engineering; Engineering","score_opus":0.0833120254126558,"score_gpt":0.30913784916280196,"score_spread":0.22582582375014615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145428772","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017534753,0.00054818945,0.49286693,0.0015731998,0.0002571987,0.00075311674,0.03239492,0.4424699,0.011601807],"genre_scores_gemma":[0.247581,0.0007525967,0.606661,0.0024990253,0.00035123332,0.0014433253,0.100085795,0.031303458,0.00932248],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99669206,0.00070465886,0.00053471443,0.00062404317,0.001202534,0.00024202354],"domain_scores_gemma":[0.9882627,0.0051478073,0.001112488,0.0031914972,0.0016582046,0.0006273083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008919743,0.0015320437,0.0015260144,0.004514329,0.0011532744,0.0039085676,0.0032789025,0.0022579762,0.017189924],"category_scores_gemma":[0.021053838,0.00092192984,0.00173875,0.004234874,0.0006558629,0.0062543266,0.0036950447,0.0014388751,0.0067285546],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009973699,0.0011831885,0.035828177,0.0035299477,0.0011867981,0.0026917937,0.0024424267,0.016226195,0.044383127,0.04815237,0.39895228,0.43545],"study_design_scores_gemma":[0.0011633626,0.00040588796,0.016145602,0.0005142892,0.00046923532,0.002359299,0.0011153934,0.2965636,0.15967642,0.055354375,0.46563148,0.0006010712],"about_ca_topic_score_codex":0.004294055,"about_ca_topic_score_gemma":0.003399448,"teacher_disagreement_score":0.017189924,"about_ca_system_score_codex":0.0016712148,"about_ca_system_score_gemma":0.0027302487,"threshold_uncertainty_score":0.057506025},"labels":[],"label_agreement":null},{"id":"W2145435867","doi":"10.1038/msb.2011.77","title":"Controlled vocabularies and semantics in systems biology","year":2011,"lang":"en","type":"article","venue":"Molecular Systems Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":331,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Terry Fox Research Institute","funders":"National Institute of General Medical Sciences; Biotechnology and Biological Sciences Research Council; European Commission","keywords":"Ontology; Terminology; Semantics (computer science); Computer science; Systems biology; Reuse; Open Biomedical Ontologies; Biology; Data science; Semantic Web; Computational biology; Information retrieval; Ontology-based data integration; Programming language; Suggested Upper Merged Ontology","score_opus":0.021414049914752778,"score_gpt":0.25758304183875635,"score_spread":0.2361689919240036,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2145435867","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014914171,0.029048732,0.8503628,0.02512681,0.0011889904,0.0007503576,0.006524012,0.0021997846,0.06988426],"genre_scores_gemma":[0.2876273,0.024465868,0.65401584,0.0057467003,0.0015830405,0.0025321352,0.012605973,0.0008310148,0.01059206],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9924326,0.0032455036,0.0016030846,0.0010507107,0.001349701,0.00031842387],"domain_scores_gemma":[0.98393434,0.010670192,0.0014300104,0.0021678284,0.0013706161,0.00042702447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010271466,0.0010407284,0.0014135177,0.0077039027,0.0035878737,0.009879626,0.0028318889,0.0032248362,0.004287957],"category_scores_gemma":[0.018170528,0.0010743728,0.0021055038,0.010433712,0.011708334,0.018468665,0.0041032326,0.0038506745,0.0012499444],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000081332055,0.0000069704993,0.00013947023,0.00016368635,0.0000148966,0.000106625244,0.0006536306,0.0013741032,0.00019348478,0.9854635,0.0025937106,0.009281829],"study_design_scores_gemma":[0.00001375813,0.000008751141,0.00015842816,0.00025739425,0.000017539516,0.00012776715,0.00038990777,0.003448258,0.00021699813,0.92667776,0.06865722,0.00002615839],"about_ca_topic_score_codex":0.014568224,"about_ca_topic_score_gemma":0.0064384956,"teacher_disagreement_score":0.014568224,"about_ca_system_score_codex":0.0045599192,"about_ca_system_score_gemma":0.006105242,"threshold_uncertainty_score":0.05432135},"labels":[],"label_agreement":null},{"id":"W2146538435","doi":"10.1186/1758-2946-3-19","title":"Linked open drug data for pharmaceutical research and development","year":2011,"lang":"en","type":"article","venue":"Journal of Cheminformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":187,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Pfizer","keywords":"Data science; Computer science; Task (project management); Open data; Data sharing; External Data Representation; Web application; World Wide Web; Data mining; Medicine; Engineering; Alternative medicine; Artificial intelligence","score_opus":0.4477610663788747,"score_gpt":0.4690782952067872,"score_spread":0.021317228827912482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2146538435","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009411902,0.009239428,0.48444423,0.02004854,0.0019681675,0.002484183,0.33370206,0.02403599,0.11466548],"genre_scores_gemma":[0.09087215,0.012795625,0.37291595,0.005327028,0.0008652868,0.0022846863,0.49768093,0.002509231,0.014749144],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98904175,0.003734181,0.0015141806,0.0010151854,0.004270122,0.00042452817],"domain_scores_gemma":[0.95214534,0.013097966,0.0050553414,0.021611976,0.0059074247,0.0021819267],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.012590253,0.0010700001,0.0013553817,0.013268317,0.001879292,0.009439941,0.0031559872,0.003020263,0.024467932],"category_scores_gemma":[0.058719076,0.0008190629,0.0020186326,0.021591099,0.0012624328,0.0098049315,0.009478928,0.0038760905,0.012049742],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007903516,0.00051335414,0.008604794,0.0039690943,0.00067814736,0.0005706637,0.00062872283,0.010027361,0.0028580218,0.28922445,0.29400378,0.3881313],"study_design_scores_gemma":[0.00012955547,0.000079524434,0.0026375027,0.00086578587,0.00015564087,0.00023950556,0.0002430727,0.007559355,0.00358627,0.13512146,0.84928924,0.00009315277],"about_ca_topic_score_codex":0.0041689593,"about_ca_topic_score_gemma":0.003732621,"teacher_disagreement_score":0.996844,"about_ca_system_score_codex":0.003116334,"about_ca_system_score_gemma":0.008772796,"threshold_uncertainty_score":0.08185345},"labels":[],"label_agreement":null},{"id":"W2147110674","doi":"10.1109/fuzzy.2010.5584301","title":"A semiotic approach to data in medical decision making","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thompson Rivers University","funders":"","keywords":"Computer science; Data science; Semiotics; Context (archaeology); Decision support system; Process (computing); Interpretation (philosophy); Representation (politics); Decision-making; Management science; Medical diagnosis; Knowledge management; Artificial intelligence; Engineering; Medicine","score_opus":0.0323315007656256,"score_gpt":0.3493943816342523,"score_spread":0.3170628808686267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147110674","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046349657,0.009636612,0.91007155,0.01654224,0.000747702,0.0002778979,0.0003288815,0.00014467632,0.057615537],"genre_scores_gemma":[0.32593155,0.009402045,0.6527546,0.002792707,0.0012675407,0.0012727482,0.0005341826,0.000106326115,0.0059383437],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97366387,0.018883534,0.00212917,0.0015227685,0.003278213,0.00052252447],"domain_scores_gemma":[0.9555922,0.036145907,0.0017924368,0.0028320358,0.0028445807,0.00079283834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018096218,0.0010259921,0.0012851984,0.007972903,0.0040671504,0.013431592,0.0024692898,0.003588709,0.0025315399],"category_scores_gemma":[0.025645042,0.00079490914,0.0017481337,0.006900079,0.032145184,0.013133862,0.0063518104,0.00572582,0.00067072787],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008934702,0.0000068151166,0.00007119301,0.000072318835,0.000007802163,0.00008222298,0.0014335923,0.00058214285,0.000089055284,0.9936732,0.00046928725,0.0035033927],"study_design_scores_gemma":[0.000010094447,0.000013087826,0.00006834132,0.00014947043,0.0000090146505,0.00015057884,0.0008211316,0.0023540133,0.00016826278,0.9687728,0.027465075,0.000018078861],"about_ca_topic_score_codex":0.0034741375,"about_ca_topic_score_gemma":0.0021328626,"teacher_disagreement_score":0.018096218,"about_ca_system_score_codex":0.0061057573,"about_ca_system_score_gemma":0.0048772823,"threshold_uncertainty_score":0.095703125},"labels":[],"label_agreement":null},{"id":"W2147320854","doi":"10.4103/0970-9290.142562","title":"PubMed alternatives to search MEDLINE: An environmental scan","year":2014,"lang":"en","type":"article","venue":"Indian Journal of Dental Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"MEDLINE; Computer science; Gateway (web page); Information retrieval; Set (abstract data type); Medicine; World Wide Web; Political science","score_opus":0.04769392886597703,"score_gpt":0.3735458151988206,"score_spread":0.32585188633284357,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147320854","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026004167,0.20726718,0.02230099,0.043709487,0.0032603845,0.016645752,0.4015382,0.009303081,0.26997074],"genre_scores_gemma":[0.087939955,0.32340002,0.23518273,0.023477904,0.0025422138,0.038224913,0.19245976,0.005836898,0.09093563],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9897503,0.0035736724,0.0039762678,0.00046707896,0.0017981647,0.00043453174],"domain_scores_gemma":[0.938652,0.03897836,0.0062273615,0.0019030037,0.012570916,0.0016683213],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.010524448,0.0014402543,0.0020922031,0.058462974,0.0013840128,0.004276328,0.0017152474,0.0019429512,0.1613162],"category_scores_gemma":[0.05355083,0.00072978117,0.0014772585,0.05748965,0.00074070075,0.0062480294,0.004434843,0.0012129211,0.04179286],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016058875,0.00015907992,0.0036872502,0.17601047,0.0005073377,0.0018490369,0.0024934495,0.00019909137,0.0036320551,0.006599018,0.2727123,0.530545],"study_design_scores_gemma":[0.00030039498,0.00024762467,0.010956155,0.05980286,0.0008171105,0.002100222,0.0029748178,0.0002780654,0.001336502,0.003950092,0.9171129,0.0001232334],"about_ca_topic_score_codex":0.0026162646,"about_ca_topic_score_gemma":0.01026458,"teacher_disagreement_score":0.99572366,"about_ca_system_score_codex":0.0019281582,"about_ca_system_score_gemma":0.009776381,"threshold_uncertainty_score":0.53965646},"labels":[],"label_agreement":null},{"id":"W2147873108","doi":"10.1093/bib/bbn052","title":"Building biomedical web communities using a semantically aware content management system","year":2008,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":38,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Computer science; World Wide Web; Content management system; Content management; Semantic Web; Information retrieval","score_opus":0.05299525054453662,"score_gpt":0.27241758063962856,"score_spread":0.21942233009509193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2147873108","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017385922,0.0003883637,0.87506145,0.0013829777,0.00020924471,0.0012706503,0.0020682213,0.09242893,0.009804139],"genre_scores_gemma":[0.07789819,0.00035709958,0.90155697,0.0005398696,0.00014699477,0.0007343255,0.0087060435,0.0035092148,0.0065513793],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9967405,0.0008104902,0.0005250257,0.00081375515,0.0009216886,0.00018850915],"domain_scores_gemma":[0.99375004,0.002038176,0.00054335373,0.0017491006,0.0012178273,0.00070140994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006069258,0.000928817,0.0010544288,0.0075804833,0.003675513,0.005848965,0.0024640805,0.0022491736,0.0044279387],"category_scores_gemma":[0.011404799,0.00093807536,0.0019554584,0.0047839344,0.0014211377,0.008780113,0.007759398,0.0019430908,0.0033610368],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009392449,0.0018429165,0.01220593,0.000996285,0.0004142517,0.002480644,0.00606048,0.025519673,0.03975349,0.10255224,0.11010766,0.6971272],"study_design_scores_gemma":[0.0002692455,0.00015553767,0.0035093082,0.00033595334,0.0002765196,0.000732186,0.0016726931,0.5473969,0.035930254,0.12605368,0.2833662,0.00030147796],"about_ca_topic_score_codex":0.00783348,"about_ca_topic_score_gemma":0.008109394,"teacher_disagreement_score":0.00783348,"about_ca_system_score_codex":0.0019581208,"about_ca_system_score_gemma":0.002992402,"threshold_uncertainty_score":0.032097697},"labels":[],"label_agreement":null},{"id":"W2148362397","doi":"10.1186/1471-2105-12-s8-s8","title":"Benchmarking of the 2010 BioCreative Challenge III text-mining competition by the BioGRID and MINT interaction databases","year":2011,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lunenfeld-Tanenbaum Research Institute; Mount Sinai Hospital","funders":"Biotechnology and Biological Sciences Research Council; National Center for Research Resources; Directorate for Biological Sciences; National Institutes of Health; Associazione Italiana per la Ricerca sul Cancro; Canadian Institutes of Health Research; European Commission","keywords":"Computer science; Annotation; Benchmarking; Identification (biology); Test set; Information retrieval; Normalization (sociology); Data curation; Named-entity recognition; Natural language processing; Set (abstract data type); Information extraction; Artificial intelligence; Database; Data mining; Task (project management)","score_opus":0.050830729813296054,"score_gpt":0.26690982688537845,"score_spread":0.2160790970720824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148362397","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38303062,0.024003832,0.020467775,0.018491987,0.004480437,0.004463146,0.45886347,0.03780372,0.048395075],"genre_scores_gemma":[0.06302171,0.0014939314,0.023666956,0.0016327957,0.00036828124,0.0011473295,0.90161794,0.001429309,0.0056217937],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.97122586,0.009572006,0.0031544648,0.0052960995,0.009194469,0.0015569843],"domain_scores_gemma":[0.96072227,0.0141088385,0.002281567,0.0051039043,0.012211679,0.0055717146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030688344,0.005638873,0.0040601934,0.0118660005,0.0037643872,0.008400519,0.009078763,0.0057506715,0.008413741],"category_scores_gemma":[0.04079245,0.0009449435,0.0037671959,0.013238704,0.0025896179,0.007244414,0.009059873,0.003399385,0.012742632],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0060496344,0.006383222,0.015059685,0.008066518,0.0020021652,0.0013375891,0.0010124499,0.02543046,0.009636577,0.0049024783,0.801294,0.11882523],"study_design_scores_gemma":[0.006600252,0.0046219337,0.08686746,0.0013571391,0.0013843023,0.0035705287,0.0038274645,0.28747514,0.032792874,0.010032779,0.5607495,0.0007206378],"about_ca_topic_score_codex":0.03898377,"about_ca_topic_score_gemma":0.04764815,"teacher_disagreement_score":0.03898377,"about_ca_system_score_codex":0.006694376,"about_ca_system_score_gemma":0.0073661776,"threshold_uncertainty_score":0.16229743},"labels":[],"label_agreement":null},{"id":"W2148866736","doi":"10.1126/science.1105776","title":"High-Throughput Mapping of a Dynamic Signaling Network in Mammalian Cells","year":2005,"lang":"en","type":"article","venue":"Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":723,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lunenfeld-Tanenbaum Research Institute; University of Toronto; Mount Sinai Hospital","funders":"National Institute of General Medical Sciences","keywords":"Interactome; Occludin; Cell biology; Transforming growth factor; Signal transduction; Transforming growth factor beta; Tight junction; Biology; Chemistry; Gene; Biochemistry","score_opus":0.01112985846536018,"score_gpt":0.2601856382676177,"score_spread":0.24905577980225752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2148866736","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38763627,0.0017914374,0.5937478,0.0004176224,0.00001839727,0.00015312144,0.008000351,0.0062771016,0.0019578564],"genre_scores_gemma":[0.5835657,0.0012638256,0.40184197,0.00006812385,0.000014000766,0.00027658392,0.01188752,0.00022447461,0.0008578247],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99958867,0.00009433879,0.000027786411,0.00013112191,0.00013904857,0.000019062312],"domain_scores_gemma":[0.9994362,0.00028701153,0.00009432827,0.00008873837,0.000070165435,0.000023467761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00043032257,0.00036922866,0.000566408,0.0015484131,0.00042870655,0.00084004365,0.00045120678,0.0003116279,0.00040163414],"category_scores_gemma":[0.0011576236,0.00022274339,0.0004868177,0.0016102192,0.00031042635,0.00079364626,0.00044008307,0.00034547248,0.0002773122],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036856264,0.00014933103,0.011281496,0.00069432764,0.00031696926,0.00057243905,0.00045838076,0.049410235,0.7865277,0.0073310365,0.0022440183,0.14064553],"study_design_scores_gemma":[0.00005412584,0.00019154428,0.04592206,0.00004559853,0.00019942886,0.00078851235,0.00041235617,0.45800716,0.4469136,0.025515964,0.021888463,0.00006123752],"about_ca_topic_score_codex":0.0017131477,"about_ca_topic_score_gemma":0.0029252544,"teacher_disagreement_score":0.0017131477,"about_ca_system_score_codex":0.0005893807,"about_ca_system_score_gemma":0.00054254394,"threshold_uncertainty_score":0.0042763352},"labels":[],"label_agreement":null},{"id":"W2149378423","doi":"10.1186/2041-1480-1-s1-s7","title":"Modeling biomedical experimental processes with OBI","year":2010,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":286,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"National Center for Research Resources; National Institute of Environmental Health Sciences; Engineering and Physical Sciences Research Council; Natural Environment Research Council; Biotechnology and Biological Sciences Research Council; Canadian Institutes of Health Research; National Institute of Mental Health; Public Health Agency; Public Health Agency of Canada; National Institutes of Health; National Institute of Biomedical Imaging and Bioengineering; Michael Smith Health Research BC; National Institute of Allergy and Infectious Diseases; NERC Environmental Bioinformatics Centre","keywords":"Computer science; Data science; Terminology; Ontology; Process (computing); Semantic Web; Interpretation (philosophy); Information retrieval; Epistemology; Programming language","score_opus":0.012008428410841534,"score_gpt":0.27673981382788215,"score_spread":0.2647313854170406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2149378423","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005085098,0.00038753534,0.98066,0.0008166676,0.00006107642,0.00042009033,0.0023029074,0.0024322204,0.00783431],"genre_scores_gemma":[0.07985296,0.0012692881,0.907733,0.0002322283,0.00008780136,0.001281773,0.0055349097,0.00047640718,0.0035314674],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9975083,0.00069710816,0.00049182156,0.00047439054,0.0006570062,0.00017133704],"domain_scores_gemma":[0.9899093,0.006176194,0.0011062666,0.001710926,0.0008542821,0.00024309102],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071722935,0.0011747243,0.0009861656,0.0054920055,0.0014282975,0.0064856634,0.0032017513,0.0017774083,0.005933884],"category_scores_gemma":[0.013664699,0.00093093805,0.0042338283,0.0053143045,0.002326279,0.0055424795,0.0044429456,0.0021115688,0.0016625315],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001565199,0.00014425766,0.0053587854,0.00089824875,0.00020428357,0.00057152385,0.0010592691,0.44974518,0.0017238441,0.46791965,0.004946759,0.06727178],"study_design_scores_gemma":[0.00003135813,0.000018854924,0.0003951303,0.00012505878,0.00012413031,0.00016081329,0.0001428058,0.76699585,0.0016112693,0.1782078,0.052161537,0.000025423762],"about_ca_topic_score_codex":0.015126308,"about_ca_topic_score_gemma":0.011217497,"teacher_disagreement_score":0.015126308,"about_ca_system_score_codex":0.0039079534,"about_ca_system_score_gemma":0.005328873,"threshold_uncertainty_score":0.037931144},"labels":[],"label_agreement":null},{"id":"W2153689380","doi":"10.3115/1567619.1567643","title":"Biomedical term recognition with the perceptron HMM algorithm","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Artificial intelligence; Hidden Markov model; Term (time); Perceptron; Multilayer perceptron; Pattern recognition (psychology); Machine learning; Identification (biology); Set (abstract data type); Feature (linguistics); Algorithm; Artificial neural network","score_opus":0.00856032494501333,"score_gpt":0.22995595804100483,"score_spread":0.2213956330959915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2153689380","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014964644,0.0009509883,0.9775965,0.0003081956,0.00016762025,0.000082189836,0.0004290039,0.004415223,0.001085706],"genre_scores_gemma":[0.27862525,0.0009210511,0.7137912,0.00031600168,0.00026337116,0.0003177314,0.0015803996,0.00016958883,0.004015469],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859244,0.00038540064,0.0001903497,0.00035293575,0.00037374374,0.00010509042],"domain_scores_gemma":[0.996977,0.001968916,0.00023687376,0.0002462994,0.0004939695,0.00007702342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022558244,0.0006972041,0.0010766576,0.0020037824,0.00040225373,0.0013203954,0.0015181914,0.0012849346,0.0019828575],"category_scores_gemma":[0.007294283,0.00042538065,0.0008745423,0.0020827744,0.00050245045,0.0020768603,0.00090140034,0.0015552887,0.0019701396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006486143,0.0002785369,0.0048539056,0.00046607954,0.00022648714,0.0003554949,0.00030508797,0.0834622,0.035630085,0.0079427995,0.008635883,0.85719484],"study_design_scores_gemma":[0.000048867794,0.00009870643,0.0017089735,0.000040735744,0.00006917809,0.0001812555,0.00003849698,0.9673214,0.013143598,0.013316235,0.0039777625,0.00005480883],"about_ca_topic_score_codex":0.0042579114,"about_ca_topic_score_gemma":0.0033233012,"teacher_disagreement_score":0.0042579114,"about_ca_system_score_codex":0.0006828018,"about_ca_system_score_gemma":0.000945188,"threshold_uncertainty_score":0.011930108},"labels":[],"label_agreement":null},{"id":"W2154264671","doi":"10.1186/1471-2164-13-s4-s10","title":"Automated extraction and semantic analysis of mutation impacts from the biomedical literature","year":2012,"lang":"en","type":"article","venue":"BMC Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ontology; Heuristics; Information retrieval; Mutation; Biomedical text mining; Information extraction; Domain (mathematical analysis); Precision and recall; Data mining; Task (project management); Stability (learning theory); Computational biology; Bioinformatics; Biology; Text mining; Genetics; Machine learning; Gene","score_opus":0.015371495529562088,"score_gpt":0.2897046117357931,"score_spread":0.274333116206231,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154264671","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31124586,0.018206174,0.3984593,0.0053969068,0.0010065355,0.002325506,0.17522724,0.06700025,0.021132262],"genre_scores_gemma":[0.29899478,0.006594196,0.5399411,0.0007492832,0.0005406503,0.0005982866,0.14827612,0.001347329,0.002958215],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99670047,0.00037957015,0.0006660844,0.00080861175,0.0013053769,0.00013981607],"domain_scores_gemma":[0.98743933,0.0060320203,0.002372977,0.00095934566,0.0028589491,0.00033736406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027542557,0.001466136,0.0009926888,0.036054797,0.0011865727,0.0027899223,0.0012230772,0.0014807315,0.0033049723],"category_scores_gemma":[0.014321108,0.00046631548,0.0017848766,0.011427406,0.00073187397,0.0035524787,0.0024182505,0.0009093657,0.002343041],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006192061,0.0005501627,0.04810779,0.010711919,0.00072023214,0.005643506,0.002249482,0.0061357496,0.096069016,0.011632575,0.074943975,0.74261636],"study_design_scores_gemma":[0.00029774164,0.0005640905,0.18854392,0.0035042702,0.002599344,0.012108079,0.004602794,0.19863251,0.16038716,0.046694204,0.38158837,0.000477494],"about_ca_topic_score_codex":0.0040675253,"about_ca_topic_score_gemma":0.0059389663,"teacher_disagreement_score":0.036054797,"about_ca_system_score_codex":0.0014380764,"about_ca_system_score_gemma":0.003548263,"threshold_uncertainty_score":0.014566064},"labels":[],"label_agreement":null},{"id":"W2154889602","doi":"10.1162/coli.2010.36.1.36101","title":"A Graph-Theoretic Framework for Semantic Distance","year":2010,"lang":"en","type":"article","venue":"Computational Linguistics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto","funders":"","keywords":"Computer science; Semantic similarity; Natural language processing; Artificial intelligence; Coherence (philosophical gambling strategy); Information retrieval; Set (abstract data type); Semantic computing; Semantic Web; Mathematics","score_opus":0.012063363554660086,"score_gpt":0.3012501952456227,"score_spread":0.2891868316909626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154889602","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001957671,0.0007313719,0.9910617,0.0006621294,0.000075972035,0.00007563854,0.0002697293,0.00012117278,0.005044563],"genre_scores_gemma":[0.14094748,0.0026087945,0.84823173,0.0005111282,0.00059446204,0.0008032148,0.0013580807,0.0002096842,0.0047354097],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9945321,0.002396101,0.00045541287,0.0011785583,0.0012695311,0.00016829664],"domain_scores_gemma":[0.99250925,0.0049747825,0.00050845277,0.00085555506,0.00095297495,0.00019910956],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005251364,0.0015311671,0.0013603845,0.00934088,0.0025188087,0.004853337,0.0034455718,0.0028235137,0.004236909],"category_scores_gemma":[0.017560605,0.00069472275,0.0021141195,0.010205042,0.0046605505,0.013107541,0.003722858,0.0035527996,0.0014075307],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012155704,0.000020981512,0.00022777682,0.00012509763,0.00004025683,0.000054216038,0.00023496979,0.01366429,0.00032011312,0.95226204,0.0017484678,0.0312895],"study_design_scores_gemma":[0.000005881496,0.000012666226,0.0001168972,0.000025750838,0.0000111034815,0.00006387986,0.000056425568,0.042079695,0.0001778888,0.94823754,0.009197041,0.000015316828],"about_ca_topic_score_codex":0.004344617,"about_ca_topic_score_gemma":0.0027010916,"teacher_disagreement_score":0.00934088,"about_ca_system_score_codex":0.0037202784,"about_ca_system_score_gemma":0.0016904869,"threshold_uncertainty_score":0.027772188},"labels":[],"label_agreement":null},{"id":"W2154996281","doi":"10.3138/jvme.34.4.431","title":"Teaching Medical Pathology in the Twenty-First Century: Virtual Microscopy Applications","year":2007,"lang":"en","type":"article","venue":"Journal of Veterinary Medical Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":61,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Virtual microscopy; Digital pathology; Curriculum; Zoom; Medical education; Histopathology; Pathology; Computer science; General pathology; Graduate medical education; Medical physics; Multimedia; Medicine; Psychology; Accreditation; Biology","score_opus":0.02343469945892704,"score_gpt":0.3728781045903753,"score_spread":0.3494434051314482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154996281","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.087610826,0.02932915,0.6350059,0.038935963,0.0020825833,0.00029216992,0.00038726957,0.004514983,0.20184122],"genre_scores_gemma":[0.29380953,0.02455416,0.64720976,0.0035027287,0.00050567643,0.00024604757,0.00047108007,0.0005413885,0.029159594],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984754,0.0007378077,0.000117291755,0.0001632501,0.00041429282,0.00009204623],"domain_scores_gemma":[0.9928635,0.003878465,0.0004390128,0.0007505873,0.0009404536,0.0011280433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044020894,0.0003759703,0.0002628496,0.0028863708,0.00086129864,0.0043103,0.0016345527,0.0013062827,0.0071855676],"category_scores_gemma":[0.007718998,0.00031447964,0.0004324198,0.00192662,0.0020725962,0.0034530347,0.004476777,0.001168098,0.0015751784],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005288372,0.00013312067,0.005852843,0.00055629166,0.000018601891,0.00048448393,0.0038117298,0.002056134,0.0043610646,0.05844452,0.019763066,0.90446526],"study_design_scores_gemma":[0.000027704778,0.0001832328,0.008102126,0.0010981817,0.000026822958,0.0030619318,0.0037874805,0.0062007816,0.004753885,0.08616613,0.88650656,0.00008521612],"about_ca_topic_score_codex":0.0018578193,"about_ca_topic_score_gemma":0.0053081126,"teacher_disagreement_score":0.0071855676,"about_ca_system_score_codex":0.0026944908,"about_ca_system_score_gemma":0.0032700128,"threshold_uncertainty_score":0.024038136},"labels":[],"label_agreement":null},{"id":"W2155261462","doi":"10.1186/2041-1480-5-46","title":"Automatically exposing OpenLifeData via SADI semantic Web Services","year":2014,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"National Bioscience Database Center; Banco Bilbao Vizcaya Argentaria; Universidad Politécnica de Madrid; Instituto Nacional de Investigación y Tecnología Agraria y Alimentaria; Fundación BBVA","keywords":"Computer science; World Wide Web; Semantic Web; Information retrieval; Data science","score_opus":0.007719603193614999,"score_gpt":0.2578920447282829,"score_spread":0.25017244153466794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155261462","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.093998544,0.00026936375,0.7441669,0.0018982633,0.00020456889,0.0016722464,0.00516446,0.1305463,0.022079471],"genre_scores_gemma":[0.38228652,0.00058860215,0.57006025,0.0013206602,0.00008075343,0.0013380594,0.022067124,0.01368532,0.008572777],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9951385,0.00076126226,0.00050176826,0.0006756732,0.0025393385,0.00038328388],"domain_scores_gemma":[0.9929724,0.0017836677,0.00052968314,0.0028852157,0.0013523399,0.00047656617],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074810274,0.0010289297,0.0006798685,0.0018251375,0.0015658605,0.004872138,0.0024353918,0.0014148975,0.0032711737],"category_scores_gemma":[0.009514969,0.0009553945,0.0016332165,0.0018604766,0.0015875704,0.005137466,0.0066025294,0.0022523573,0.0022383535],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032840923,0.0026420772,0.0826508,0.0019176651,0.0005174302,0.0049760235,0.012575063,0.041729186,0.1407484,0.19833781,0.1086544,0.40196702],"study_design_scores_gemma":[0.00028679386,0.0001615471,0.007406431,0.0002669017,0.00016275948,0.0015022826,0.0019894459,0.36083975,0.20774955,0.09658211,0.3227538,0.0002986254],"about_ca_topic_score_codex":0.0054801474,"about_ca_topic_score_gemma":0.005039909,"teacher_disagreement_score":0.0074810274,"about_ca_system_score_codex":0.0029431193,"about_ca_system_score_gemma":0.0033192902,"threshold_uncertainty_score":0.039563894},"labels":[],"label_agreement":null},{"id":"W2155313075","doi":"10.1109/wi.2007.37","title":"Automatic Taxonomy Extraction Using Google and Term Dependency","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Terminology; Dependency (UML); Taxonomy (biology); Data mining; Adjacency matrix; Information extraction; Information retrieval; Term (time); Adjacency list; Formal concept analysis; Artificial intelligence; Theoretical computer science; Algorithm","score_opus":0.03794426146626626,"score_gpt":0.3148330309274555,"score_spread":0.27688876946118923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155313075","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05040755,0.0025896537,0.89460003,0.00072264974,0.00024530824,0.0011891403,0.01968816,0.02171772,0.008839798],"genre_scores_gemma":[0.082479,0.0011876913,0.88505965,0.00008473905,0.00006228727,0.00070357375,0.027167916,0.00044949833,0.0028056446],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99815875,0.0002170865,0.0002211831,0.00031230034,0.00096778944,0.00012296588],"domain_scores_gemma":[0.99735045,0.0008089057,0.0003095267,0.0002496246,0.0011831862,0.00009829633],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000975199,0.0013150839,0.0014158012,0.025021223,0.001648742,0.001837954,0.0011136694,0.0009594218,0.0025478906],"category_scores_gemma":[0.0065018698,0.0006330488,0.001671366,0.017199239,0.00032653165,0.0037186567,0.001877152,0.0007866173,0.002598213],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001941259,0.00019989617,0.0069158017,0.0015888391,0.00026258721,0.0009337681,0.0009187001,0.007005883,0.030238396,0.023736399,0.049564984,0.8784406],"study_design_scores_gemma":[0.0001292779,0.0002187693,0.026285809,0.0007323401,0.00052135973,0.004334554,0.0018684752,0.5474687,0.061268847,0.10935771,0.24742204,0.00039206853],"about_ca_topic_score_codex":0.009261054,"about_ca_topic_score_gemma":0.016075473,"teacher_disagreement_score":0.025021223,"about_ca_system_score_codex":0.0012347355,"about_ca_system_score_gemma":0.0035536587,"threshold_uncertainty_score":0.018414259},"labels":[],"label_agreement":null},{"id":"W2155361294","doi":"10.3115/1572364.1572371","title":"Extraction of named entities from tables in gene mutation literature","year":2009,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Australian Research Council; Memorial University of Newfoundland; Australian Government","keywords":"Computer science; Named entity; Task (project management); Information extraction; Mutation; Feature (linguistics); Artificial intelligence; Information retrieval; Named-entity recognition; Natural language processing; Data mining; Gene; Genetics; Biology; Engineering","score_opus":0.007560293768626997,"score_gpt":0.2658806666438477,"score_spread":0.2583203728752207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155361294","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.20798622,0.016590577,0.50992435,0.0037659425,0.00090682984,0.0012906695,0.22953406,0.017946733,0.012054653],"genre_scores_gemma":[0.24555995,0.0068030194,0.5703584,0.0004021636,0.00037784188,0.00045445032,0.17323777,0.0005825724,0.0022239268],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99809533,0.0003651507,0.000498425,0.00040056463,0.0005525697,0.00008801525],"domain_scores_gemma":[0.97729355,0.015047819,0.003214717,0.0016287399,0.0024462712,0.0003688474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003101495,0.0010242654,0.0012968897,0.023080574,0.0010206826,0.002744826,0.0011589002,0.0011346375,0.0038748942],"category_scores_gemma":[0.019479806,0.00042269338,0.0011765738,0.021280503,0.00043518085,0.0038005041,0.0012133232,0.0006737136,0.002030037],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008229746,0.00041320405,0.05535159,0.00866576,0.0006361093,0.008226581,0.002162319,0.02014538,0.0613929,0.02819165,0.04912371,0.7648678],"study_design_scores_gemma":[0.00035815552,0.0007824119,0.093868256,0.0037850863,0.0020527176,0.013878723,0.004184067,0.20213984,0.18707941,0.09483597,0.39656705,0.0004683544],"about_ca_topic_score_codex":0.0024251218,"about_ca_topic_score_gemma":0.0036258365,"teacher_disagreement_score":0.023080574,"about_ca_system_score_codex":0.00091971003,"about_ca_system_score_gemma":0.002204124,"threshold_uncertainty_score":0.016402483},"labels":[],"label_agreement":null},{"id":"W2155744163","doi":"10.1093/bib/6.3.222","title":"Hairpins in bookstacks: Information retrieval from biomedical text","year":2005,"lang":"en","type":"review","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Biomedicine; Computer science; Information retrieval; Field (mathematics); Identification (biology); Data curation; Data science; Bioinformatics; Biology","score_opus":0.027332834017884496,"score_gpt":0.3142867120808825,"score_spread":0.286953878062998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2155744163","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018821319,0.89947015,0.04892564,0.0061054626,0.0036736298,0.00030578097,0.0012326452,0.0025799067,0.03582464],"genre_scores_gemma":[0.011673217,0.79139626,0.10132633,0.008910062,0.0040962403,0.0003516117,0.005366705,0.00068772933,0.07619185],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99874276,0.00025699864,0.00015700009,0.00017325224,0.0006096798,0.000060330043],"domain_scores_gemma":[0.99766505,0.0012467677,0.0001957932,0.00019559557,0.000554736,0.00014211383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014776932,0.001106028,0.002181567,0.009674345,0.0008246943,0.0039226837,0.0023617158,0.0023048643,0.024761487],"category_scores_gemma":[0.005408972,0.00052886084,0.00096456736,0.013154811,0.0016522156,0.008539993,0.001962373,0.0015805381,0.032425545],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040849547,0.000044238703,0.00018709038,0.003991541,0.00006899481,0.00015102541,0.00027155058,0.00025767062,0.0017778741,0.006446129,0.12574697,0.8610161],"study_design_scores_gemma":[0.00002559752,0.00005080635,0.0008627962,0.0015773107,0.000066521425,0.0016034812,0.00027029938,0.0005522174,0.0024570785,0.012107123,0.98036295,0.00006370333],"about_ca_topic_score_codex":0.0017260045,"about_ca_topic_score_gemma":0.0024692207,"teacher_disagreement_score":0.024761487,"about_ca_system_score_codex":0.0008698952,"about_ca_system_score_gemma":0.0013362716,"threshold_uncertainty_score":0.082835436},"labels":[],"label_agreement":null},{"id":"W2156124616","doi":"10.1093/bioinformatics/btm452","title":"Leveraging the structure of the Semantic Web to enhance information retrieval for proteomics","year":2007,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; National Institute of General Medical Sciences; McGill University","keywords":"Computer science; Information retrieval; Graph; Leverage (statistics); SPARQL; RDF; Subgraph isomorphism problem; Semantic Web; Artificial intelligence; Theoretical computer science","score_opus":0.009301903491705083,"score_gpt":0.26632088811460186,"score_spread":0.2570189846228968,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2156124616","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15303357,0.005477612,0.809536,0.005016966,0.00027485302,0.000940699,0.0034504174,0.012366765,0.009903121],"genre_scores_gemma":[0.2867178,0.002414681,0.7035959,0.0006377786,0.0002010319,0.00025868797,0.0044201445,0.00045169843,0.0013023445],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997387,0.0009877756,0.00029297415,0.00029260528,0.00092268497,0.0001169953],"domain_scores_gemma":[0.9920372,0.004657482,0.00056945597,0.0011660288,0.0014364081,0.0001334738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037535478,0.0008832243,0.00084979774,0.009231799,0.00074400136,0.0025207659,0.0008015461,0.0011279554,0.001760841],"category_scores_gemma":[0.018169725,0.00042673864,0.0010824885,0.007927314,0.0008529571,0.007457106,0.0023140989,0.00090730883,0.0018901817],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005471095,0.0006863999,0.010222214,0.0021760177,0.00025719285,0.00036177252,0.0007189561,0.04106189,0.07408066,0.028442556,0.017180905,0.8242642],"study_design_scores_gemma":[0.00019606792,0.0006516267,0.008238292,0.00032640548,0.00034590272,0.0010978717,0.00072240364,0.6542586,0.10126102,0.17737913,0.05531592,0.00020670658],"about_ca_topic_score_codex":0.0032391238,"about_ca_topic_score_gemma":0.0058002654,"teacher_disagreement_score":0.009231799,"about_ca_system_score_codex":0.0010083315,"about_ca_system_score_gemma":0.0014817577,"threshold_uncertainty_score":0.01985085},"labels":[],"label_agreement":null},{"id":"W2157414955","doi":"10.1186/2041-1480-5-5","title":"BioHackathon series in 2011 and 2012: penetration of ontology and linked data in life science domains","year":2014,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University; Ontario Institute for Cancer Research","funders":"National Bioscience Database Center; National Institute of General Medical Sciences; Biotechnology and Biological Sciences Research Council; Ministry of Economy, Trade and Industry; Ministry of Education, Culture, Sports, Science and Technology","keywords":"Computer science; Interoperability; Ontology; RDF; Semantic Web; Metadata; Semantic interoperability; World Wide Web; Data science; Linked data; Visualization; Semantic integration; Semantic technology; Information retrieval; Semantic Web Stack; Data mining","score_opus":0.021279614003586723,"score_gpt":0.29392193949173573,"score_spread":0.272642325488149,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157414955","genre_codex":"editorial","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14523433,0.22555698,0.045470476,0.17777257,0.22612351,0.0020495506,0.009850451,0.0014144024,0.16652767],"genre_scores_gemma":[0.26429343,0.1605581,0.052289426,0.027933713,0.046049032,0.0023842677,0.021894228,0.0020702663,0.42252752],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963464,0.0006592688,0.00019616779,0.0005265145,0.0016027929,0.00066872797],"domain_scores_gemma":[0.98474914,0.0038347535,0.0012641123,0.00081201276,0.004186349,0.005153674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011364479,0.0009793708,0.00066326704,0.005260099,0.0028467986,0.004103131,0.0012106352,0.0025510278,0.009943241],"category_scores_gemma":[0.010608578,0.00042007878,0.0006775948,0.0065137055,0.0018803278,0.0052621122,0.005860133,0.0041205916,0.0019964837],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005107713,0.00032455145,0.0036936712,0.001336139,0.000062268664,0.0015733674,0.0065199407,0.0006861335,0.0139462035,0.03219241,0.5937029,0.34545162],"study_design_scores_gemma":[0.000013501482,0.00014789647,0.007853748,0.00023424476,0.000015180832,0.00047769115,0.0019949444,0.00023912922,0.0032611112,0.0015535511,0.98415595,0.000052968782],"about_ca_topic_score_codex":0.0050069285,"about_ca_topic_score_gemma":0.011009173,"teacher_disagreement_score":0.011364479,"about_ca_system_score_codex":0.0041754814,"about_ca_system_score_gemma":0.007533478,"threshold_uncertainty_score":0.060101807},"labels":[],"label_agreement":null},{"id":"W2157423684","doi":"10.1007/bf03020002","title":"A system of classification for the clinical uses of capnography","year":2007,"lang":"en","type":"article","venue":"Canadian Journal of Anesthesia/Journal canadien d anesthésie","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sunnybrook Hospital; Health Sciences Centre; Sunnybrook Health Science Centre","funders":"","keywords":"Capnography; Medicine; Computer science; Anesthesia","score_opus":0.03485068520784841,"score_gpt":0.29001181259053094,"score_spread":0.25516112738268254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157423684","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09100706,0.002898108,0.77863854,0.0070334724,0.0010482247,0.0037488162,0.05962321,0.026642675,0.029359875],"genre_scores_gemma":[0.17612801,0.0010651273,0.7817165,0.00062206323,0.0001638261,0.0013599308,0.033198405,0.00044183765,0.0053043202],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.996682,0.00044495479,0.00080999016,0.00084478175,0.000985797,0.00023246768],"domain_scores_gemma":[0.9940363,0.0022488914,0.0005038066,0.00055353966,0.0023246445,0.0003327282],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031464652,0.001001357,0.00082091056,0.010091301,0.0019995607,0.0035183495,0.0015613498,0.001915187,0.005625302],"category_scores_gemma":[0.012403384,0.00035726975,0.0019737997,0.0054715495,0.0007875084,0.0028676046,0.0018922357,0.0013370501,0.004202651],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075386086,0.0005425526,0.0690005,0.001134134,0.000477654,0.0010012712,0.0017716156,0.0041600205,0.011241335,0.028782427,0.07444815,0.8066865],"study_design_scores_gemma":[0.00042144142,0.00087654893,0.12776642,0.0026379141,0.0026248223,0.0057400023,0.0033694003,0.2703122,0.037895914,0.11985378,0.42802066,0.00048087107],"about_ca_topic_score_codex":0.0148040205,"about_ca_topic_score_gemma":0.014543068,"teacher_disagreement_score":0.0148040205,"about_ca_system_score_codex":0.00239985,"about_ca_system_score_gemma":0.0053745895,"threshold_uncertainty_score":0.029435694},"labels":[],"label_agreement":null},{"id":"W2157789609","doi":"10.1002/meet.1450390131","title":"Filtering for medical news items","year":2002,"lang":"en","type":"article","venue":"Proceedings of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Context (archaeology); Audience measurement; Filter (signal processing); Task (project management); Process (computing); Service (business); Information retrieval; Medical information; World Wide Web; Advertising","score_opus":0.013888660793125618,"score_gpt":0.2678104120853601,"score_spread":0.2539217512922345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2157789609","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3677821,0.03576891,0.41088635,0.016401099,0.008083201,0.005826693,0.07609968,0.019655794,0.059496157],"genre_scores_gemma":[0.3669469,0.01543001,0.43528587,0.0044636666,0.008934984,0.0010497889,0.12436812,0.0014368541,0.0420838],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99380916,0.0010678957,0.0010386492,0.0008202248,0.0029190052,0.00034510597],"domain_scores_gemma":[0.95637083,0.022491682,0.0047478746,0.0032394126,0.011798501,0.001351735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046522096,0.0011339281,0.0016562886,0.029289113,0.002024842,0.0053697177,0.0016085597,0.0019807308,0.011656086],"category_scores_gemma":[0.030089328,0.00059036934,0.0014218654,0.012537438,0.00057626644,0.002855596,0.0013125782,0.0012267318,0.009284609],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011367423,0.00041899487,0.052806303,0.003588229,0.00047997152,0.0027520068,0.0010494206,0.0012245742,0.015999414,0.0033934184,0.11379392,0.80335706],"study_design_scores_gemma":[0.00035542698,0.00092320255,0.19274291,0.0021613266,0.0021224972,0.0123982,0.0042744996,0.046872627,0.09469769,0.021948,0.6211164,0.00038725775],"about_ca_topic_score_codex":0.0050577447,"about_ca_topic_score_gemma":0.0055577704,"teacher_disagreement_score":0.029289113,"about_ca_system_score_codex":0.0009489585,"about_ca_system_score_gemma":0.0014406526,"threshold_uncertainty_score":0.038993478},"labels":[],"label_agreement":null},{"id":"W2158040974","doi":"10.1016/s1386-5056(02)00050-3","title":"Getting to the (c)ore of knowledge: mining biomedical literature","year":2002,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":125,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Biomedical text mining; Computer science; Categorization; Data science; Reading (process); Knowledge extraction; The Internet; Domain (mathematical analysis); Information extraction; Process (computing); Information retrieval; Scientific literature; World Wide Web; Text mining; Natural language processing; Data mining; Artificial intelligence; Linguistics","score_opus":0.01995399975705435,"score_gpt":0.31765983469779346,"score_spread":0.2977058349407391,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158040974","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2251253,0.034486875,0.6535662,0.029795578,0.0017356086,0.001993471,0.026740618,0.0052262656,0.021330148],"genre_scores_gemma":[0.22635442,0.010146675,0.7404864,0.0022818274,0.0006848551,0.00046342443,0.015839338,0.00036660524,0.0033764995],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968035,0.0005578497,0.00062951585,0.00064856396,0.0012049554,0.00015554272],"domain_scores_gemma":[0.99019086,0.0055812257,0.0010902901,0.00093038613,0.0018106261,0.00039660014],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023838514,0.0009415828,0.0014081877,0.019242609,0.0017108495,0.0050102985,0.0017017437,0.0020832613,0.0029807154],"category_scores_gemma":[0.024865894,0.00048946106,0.0018398279,0.013353289,0.0018464484,0.0070923176,0.0044297785,0.0016891307,0.002022038],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033983504,0.00040411443,0.017168282,0.0055510444,0.00060238177,0.0026761342,0.003937211,0.0026301246,0.015036282,0.024319788,0.039902084,0.88743275],"study_design_scores_gemma":[0.00032309757,0.0006434132,0.032603342,0.0069492077,0.0036766091,0.0092777815,0.021085203,0.07895931,0.03532881,0.42365894,0.38708562,0.00040855067],"about_ca_topic_score_codex":0.004051882,"about_ca_topic_score_gemma":0.0070289248,"teacher_disagreement_score":0.019242609,"about_ca_system_score_codex":0.0009454907,"about_ca_system_score_gemma":0.0046524527,"threshold_uncertainty_score":0.012607217},"labels":[],"label_agreement":null},{"id":"W2158351585","doi":"10.1093/bib/bbn051","title":"Moby and Moby 2: Creatures of the Deep (Web)","year":2009,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital","funders":"Natural Sciences and Engineering Research Council of Canada; Genome Alberta; Microsoft Research; Heart and Stroke Foundation of British Columbia and Yukon; Canadian Institutes of Health Research; Genome Canada; Heart and Stroke Foundation of Canada","keywords":"Computer science; Web standards; Semantic Web; World Wide Web; Social Semantic Web; OWL-S; Semantic Web Stack; Data Web; Semantic analytics; RDF; Linked data; The Internet; Web service","score_opus":0.005709647186809341,"score_gpt":0.23024597148705558,"score_spread":0.22453632430024625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158351585","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06680907,0.020024728,0.5553895,0.16199343,0.0048883082,0.00047664964,0.001214796,0.005774343,0.18342912],"genre_scores_gemma":[0.30431882,0.010731968,0.50666577,0.04079393,0.0015769633,0.0006034857,0.0017656639,0.0033710937,0.1301724],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99675995,0.0012376215,0.0001139791,0.0005105664,0.0010333139,0.00034467044],"domain_scores_gemma":[0.9942532,0.00230429,0.00037793015,0.0016978093,0.0006206613,0.00074601127],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008376222,0.00063377776,0.00047208613,0.0020245402,0.0036876919,0.010491695,0.0013674032,0.0036126766,0.0050670058],"category_scores_gemma":[0.011068115,0.00061363354,0.00071122224,0.002032449,0.013123637,0.023720127,0.010052573,0.004668011,0.0015759247],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000108276494,0.000041809883,0.0019761755,0.00018581931,0.000036474674,0.0003058464,0.019859402,0.00059101573,0.0031533404,0.7642786,0.0628199,0.14664328],"study_design_scores_gemma":[0.00001125649,0.000024711848,0.0008089095,0.0001854844,0.000012674179,0.00031102303,0.0031407045,0.0014464805,0.0011790533,0.16974685,0.82309395,0.00003890579],"about_ca_topic_score_codex":0.006994473,"about_ca_topic_score_gemma":0.0061713196,"teacher_disagreement_score":0.010491695,"about_ca_system_score_codex":0.0015846745,"about_ca_system_score_gemma":0.0032512052,"threshold_uncertainty_score":0.04429823},"labels":[],"label_agreement":null},{"id":"W2158439765","doi":"10.1186/1741-7007-5-44","title":"Broadening the horizon – level 2.5 of the HUPO-PSI format for molecular interactions","year":2007,"lang":"en","type":"article","venue":"BMC Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":276,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Sinai Hospital; University of Toronto; Lunenfeld-Tanenbaum Research Institute","funders":"National Institutes of Health; European Commission; Genome Canada; Ontario Genomics; National Institute of General Medical Sciences; Ontario Genomics Institute; U.S. Department of Energy","keywords":"Collaboratory; Human proteome project; Computer science; Schema (genetic algorithms); XML; Data science; World Wide Web; Computational biology; Biology; Information retrieval; Proteomics","score_opus":0.08892370149535112,"score_gpt":0.35643402665071994,"score_spread":0.2675103251553688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2158439765","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008657208,0.0014732833,0.76017606,0.0057432964,0.0024627128,0.0012156054,0.05534709,0.11711761,0.047807094],"genre_scores_gemma":[0.07009438,0.0026728837,0.62698567,0.0075901514,0.0014469407,0.0020045368,0.23421085,0.03292615,0.022068465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9885781,0.0027255751,0.0032413988,0.0011187935,0.003380301,0.0009557729],"domain_scores_gemma":[0.97128594,0.007887597,0.0019703784,0.012444528,0.0053822254,0.0010293971],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018816527,0.0022007956,0.0014786883,0.005096599,0.0017598976,0.01148014,0.005114872,0.0042532254,0.028205141],"category_scores_gemma":[0.037110757,0.00225139,0.0029923536,0.0047291727,0.0020483315,0.016167436,0.009868508,0.008260399,0.03080922],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014945994,0.0003518412,0.0058049704,0.0023604883,0.000203595,0.0010745559,0.0024123383,0.006000102,0.015005965,0.38722014,0.40166286,0.17640859],"study_design_scores_gemma":[0.00009166381,0.00011297939,0.0012134094,0.00085327285,0.00007153328,0.0005165684,0.00024119492,0.010407167,0.0129586505,0.06278113,0.9105724,0.00018018347],"about_ca_topic_score_codex":0.005435786,"about_ca_topic_score_gemma":0.0023741934,"teacher_disagreement_score":0.028205141,"about_ca_system_score_codex":0.002776628,"about_ca_system_score_gemma":0.006119358,"threshold_uncertainty_score":0.09951252},"labels":[],"label_agreement":null},{"id":"W2159230276","doi":"10.1186/1471-2105-9-s11-s10","title":"Recognizing speculative language in biomedical research articles: a linguistically motivated perspective","year":2008,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":126,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Collège de Maisonneuve; Concordia University","funders":"","keywords":"Weighting; Computer science; Perspective (graphical); Sentence; Artificial intelligence; Natural language processing; Point (geometry); Scheme (mathematics); Machine learning; Mathematics","score_opus":0.09097231171662291,"score_gpt":0.36817676912295266,"score_spread":0.27720445740632976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159230276","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2928649,0.006336572,0.6434589,0.0070212707,0.00067832216,0.0011413686,0.008060624,0.025651952,0.014786094],"genre_scores_gemma":[0.42277995,0.001250708,0.5645206,0.0010723894,0.00053101406,0.0002724692,0.006495194,0.0004623568,0.0026153673],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9964045,0.0010042025,0.0005649048,0.0008120368,0.0010745046,0.00013981976],"domain_scores_gemma":[0.9785943,0.012075517,0.004012039,0.0018370869,0.002965168,0.00051581446],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0047557056,0.0008938216,0.00069762074,0.007906768,0.0006308471,0.0039170873,0.0013605853,0.0016156989,0.0022470066],"category_scores_gemma":[0.022767514,0.00040861944,0.0010322512,0.002808597,0.00088528893,0.005000348,0.002036643,0.0017167676,0.0022327534],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009945789,0.00044273285,0.025972921,0.0029125041,0.0003876229,0.0015171317,0.00411373,0.0052586338,0.18408361,0.013791716,0.021945909,0.7385789],"study_design_scores_gemma":[0.00019953446,0.0012335074,0.069585726,0.0016572108,0.0010030468,0.006475314,0.00608251,0.43663132,0.23429804,0.10330154,0.1389236,0.00060870306],"about_ca_topic_score_codex":0.0010295082,"about_ca_topic_score_gemma":0.0021221833,"teacher_disagreement_score":0.99524426,"about_ca_system_score_codex":0.00084704824,"about_ca_system_score_gemma":0.001224869,"threshold_uncertainty_score":0.025150895},"labels":[],"label_agreement":null},{"id":"W2159856811","doi":"10.2196/medinform.3531","title":"Design and Development of a Linked Open Data-Based Health Information Representation and Visualization System: Potentials and Preliminary Evaluation","year":2014,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Visualization; Computer science; Information visualization; Representation (politics); Data visualization; Data science; Human–computer interaction; Data mining","score_opus":0.08560215020500526,"score_gpt":0.40013744905921816,"score_spread":0.3145352988542129,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2159856811","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04780238,0.00032280062,0.8987632,0.0017349994,0.00021818613,0.0063328184,0.0032580341,0.03560043,0.005967099],"genre_scores_gemma":[0.06919758,0.000280926,0.9169635,0.00040018134,0.000028481201,0.0023746812,0.006569702,0.0011997113,0.0029852828],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962798,0.0014232033,0.00037095137,0.0005997263,0.0010987781,0.00022759137],"domain_scores_gemma":[0.9944688,0.0019387684,0.00026117783,0.0008256051,0.0019068267,0.0005987735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009874805,0.0008730499,0.0007112527,0.0014786763,0.0009509234,0.003707984,0.0032558613,0.0018021112,0.0076170815],"category_scores_gemma":[0.013116419,0.0006604094,0.0014249763,0.0012455703,0.0009232673,0.0044864058,0.002823394,0.0016996914,0.0019319087],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004094591,0.0035585144,0.020932,0.0047797128,0.00084616087,0.0023433205,0.007358626,0.058601443,0.08728099,0.039547574,0.0531474,0.7175097],"study_design_scores_gemma":[0.0017263213,0.003248931,0.0100662075,0.0010153186,0.0008467147,0.0009319745,0.0022281385,0.63180137,0.10154278,0.0211591,0.22499071,0.00044246728],"about_ca_topic_score_codex":0.00683057,"about_ca_topic_score_gemma":0.0046611787,"teacher_disagreement_score":0.009874805,"about_ca_system_score_codex":0.0015469166,"about_ca_system_score_gemma":0.0035921396,"threshold_uncertainty_score":0.052223563},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"bench_or_experimental","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"not_applicable","genre":"software","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W2160014933","doi":"10.1016/j.jbi.2005.02.001","title":"Identifying reasoning strategies in medical decision making: A methodological guide","year":2005,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":115,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Identification (biology); Context (archaeology); Process (computing); Frame (networking); Management science; Data science; Artificial intelligence; Knowledge management","score_opus":0.14896582713966713,"score_gpt":0.48605511682021446,"score_spread":0.3370892896805473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160014933","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00071294676,0.026618883,0.954032,0.009256566,0.00021771317,0.0013498556,0.0006913413,0.00047842553,0.0066422205],"genre_scores_gemma":[0.005336648,0.01297718,0.9782602,0.0006903018,0.000098044555,0.0009855453,0.00052240014,0.0000547744,0.0010747908],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97944176,0.010730504,0.0042602574,0.0017124155,0.003472152,0.00038293606],"domain_scores_gemma":[0.95834374,0.033106774,0.0014429977,0.0021883429,0.0044491463,0.00046901623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03846546,0.0036621878,0.0038935,0.013334192,0.0025262907,0.015887773,0.010926145,0.004704026,0.005972036],"category_scores_gemma":[0.03895482,0.002529921,0.0046169753,0.009222567,0.009522499,0.014875265,0.0056385146,0.008395362,0.0039464133],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000083458785,0.00025110148,0.00096763164,0.007597306,0.00043516897,0.00047610287,0.0029936922,0.0050171185,0.0013564157,0.57020766,0.018222375,0.39239195],"study_design_scores_gemma":[0.00008238141,0.000049155795,0.00037573124,0.005407736,0.00025432123,0.00066139386,0.0022824323,0.0121347755,0.0017211591,0.8251116,0.15183191,0.00008728221],"about_ca_topic_score_codex":0.007100502,"about_ca_topic_score_gemma":0.008320684,"teacher_disagreement_score":0.03846546,"about_ca_system_score_codex":0.005187928,"about_ca_system_score_gemma":0.013688763,"threshold_uncertainty_score":0.2034272},"labels":[],"label_agreement":null},{"id":"W2160476655","doi":"10.1177/1460458208096556","title":"Topic maps for exploring nosological, lexical, semantic and HL7 structures for clinical data","year":2008,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke; Dalhousie University","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Information retrieval; Semantic interoperability; Semantic integration; Referent; Terminology; Systematized Nomenclature of Medicine; Natural language processing; SNOMED CT; Interoperability; Linguistics; Semantic Web; World Wide Web; Semantic computing","score_opus":0.4838173379846424,"score_gpt":0.4623620342121831,"score_spread":0.0214553037724593,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160476655","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016254898,0.0002809856,0.9633649,0.00044212284,0.000055964203,0.0004927027,0.005357751,0.00975166,0.0039990665],"genre_scores_gemma":[0.1010856,0.00027806885,0.8866003,0.0000641288,0.000034737146,0.0007206871,0.008433418,0.000940375,0.0018426791],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99798465,0.0007887727,0.00027628255,0.00038994348,0.00044755233,0.00011280898],"domain_scores_gemma":[0.99070835,0.006926225,0.0004681169,0.001018662,0.00063487946,0.00024371453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057577677,0.0008845171,0.0005292562,0.006105228,0.0015814992,0.005017797,0.0010729418,0.001107904,0.0074481573],"category_scores_gemma":[0.019393936,0.00075298856,0.0023104828,0.006699705,0.00101514,0.007991638,0.0041394667,0.0012553278,0.0016792902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00080767047,0.00027811425,0.021751272,0.0022039067,0.00042910312,0.0017007173,0.02143249,0.023133008,0.01041167,0.2431461,0.031226896,0.643479],"study_design_scores_gemma":[0.00016501576,0.0003497761,0.013756543,0.0009790572,0.00043263848,0.0016423729,0.010253266,0.3330861,0.016851034,0.3837649,0.23845826,0.00026100164],"about_ca_topic_score_codex":0.0064195413,"about_ca_topic_score_gemma":0.009659842,"teacher_disagreement_score":0.0074481573,"about_ca_system_score_codex":0.0015745418,"about_ca_system_score_gemma":0.0022662103,"threshold_uncertainty_score":0.030450344},"labels":[],"label_agreement":null},{"id":"W2160615485","doi":"10.1093/bioinformatics/bti120","title":"A Java tool for dynamic web-based 3D visualization of anatomy and overlapping gene or protein expression patterns","year":2004,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Java; Visualization; Computer science; Computational biology; Expression (computer science); Web application; Gene expression; Computer graphics (images); Gene; Biology; Anatomy; Programming language; World Wide Web; Data mining; Genetics","score_opus":0.011974473710474309,"score_gpt":0.2844673868023079,"score_spread":0.27249291309183354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160615485","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009726354,0.00019181025,0.6925007,0.00016676242,0.00012336342,0.00025365586,0.014810496,0.28590643,0.005074177],"genre_scores_gemma":[0.017261004,0.0008098364,0.81694704,0.0006834875,0.00012566008,0.0019572126,0.057289183,0.084035955,0.020890737],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99923456,0.000085918604,0.000108409804,0.00015636253,0.0003430331,0.00007172219],"domain_scores_gemma":[0.9983859,0.00079791644,0.00012779732,0.0002506014,0.00023157023,0.00020619099],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017236859,0.001996765,0.001447535,0.0031949559,0.00070249534,0.0024557854,0.003265068,0.0014111828,0.08236214],"category_scores_gemma":[0.0030043942,0.0018393332,0.0016742646,0.0016733268,0.0005786491,0.0023528545,0.003597781,0.0024287354,0.040066384],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053670024,0.00038486664,0.0018593312,0.0019424146,0.0002552904,0.0010850724,0.00044952816,0.004484866,0.05592012,0.0113702,0.5138743,0.4078374],"study_design_scores_gemma":[0.0007512148,0.00014595373,0.00679949,0.00056717324,0.00020218903,0.002770056,0.0001538154,0.10151542,0.054981228,0.033090264,0.7986494,0.0003738536],"about_ca_topic_score_codex":0.0024351834,"about_ca_topic_score_gemma":0.004203144,"teacher_disagreement_score":0.08236214,"about_ca_system_score_codex":0.00054577837,"about_ca_system_score_gemma":0.0011307175,"threshold_uncertainty_score":0.27552885},"labels":[],"label_agreement":null},{"id":"W2160653801","doi":"10.1503/cmaj.080710","title":"Canada's pathology","year":2008,"lang":"en","type":"editorial","venue":"Canadian Medical Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Pathology; Forensic pathology; Anatomical pathology; Breast cancer; Medicine; Molecular pathology; Computer science; Surgical pathology; Data science; Cancer; Biology; Internal medicine","score_opus":0.004015937033936851,"score_gpt":0.22106301966529587,"score_spread":0.21704708263135902,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2160653801","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008734552,0.011249409,0.00019051578,0.21686047,0.7593265,0.000043017637,0.0003418776,0.000119319215,0.011781462],"genre_scores_gemma":[0.0038420244,0.027161015,0.0008593227,0.19861464,0.6841342,0.00006635042,0.0003895145,0.00014481503,0.08478809],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9919274,0.0004925756,0.00061750866,0.000607373,0.0057471115,0.00060800987],"domain_scores_gemma":[0.96165127,0.0063796705,0.00072633795,0.000969578,0.024847629,0.005425526],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0063068205,0.0019391702,0.0015895761,0.004806166,0.008318418,0.008898165,0.003940901,0.018162042,0.01589957],"category_scores_gemma":[0.025089193,0.0010545201,0.0016569654,0.003201543,0.0062724636,0.0029590055,0.0015610028,0.020980235,0.006809012],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000026813555,0.0000012790084,0.00001553248,0.000033118296,0.0000023408506,0.000035144403,0.000009475257,0.000008095711,0.000009839084,0.00025387775,0.997442,0.0021864977],"study_design_scores_gemma":[0.0000067031456,0.0000024737221,0.0002044888,0.00013464809,0.000009183213,0.00006959205,0.00004116871,0.000025239666,0.000019309451,0.00029543543,0.9991837,0.000007965338],"about_ca_topic_score_codex":0.56181026,"about_ca_topic_score_gemma":0.8136637,"teacher_disagreement_score":0.43818974,"about_ca_system_score_codex":0.03612032,"about_ca_system_score_gemma":0.08340903,"threshold_uncertainty_score":0.88154066},"labels":[],"label_agreement":null},{"id":"W2161088626","doi":"10.1016/j.febslet.2012.12.013","title":"Corrigendum to “From phenotype to gene: Detecting disease‐specific gene functional modules via a text‐based human disease phenotype network construction” [FEBS Lett. 584 (2010) 3635–3643]","year":2012,"lang":"en","type":"erratum","venue":"FEBS Letters","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Phenotype; Gene; Clinical phenotype; Genetics; Disease; Computational biology; Biology; Medicine; Pathology","score_opus":0.027956824199573387,"score_gpt":0.23661639285519714,"score_spread":0.20865956865562377,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161088626","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0001365176,0.0017097477,0.001921389,0.047963154,0.9400292,0.00003864374,0.0023390653,0.00090546143,0.004956857],"genre_scores_gemma":[0.008749472,0.010921749,0.009717233,0.12859178,0.262984,0.00034352957,0.0143327415,0.0034204738,0.560939],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99646807,0.00056659704,0.0005571894,0.00052087946,0.0015905093,0.00029680232],"domain_scores_gemma":[0.97865033,0.0051042796,0.0009905843,0.0011658341,0.01315098,0.00093798625],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026471775,0.0027678802,0.0025481125,0.00480495,0.0032039555,0.003724802,0.004088188,0.005085214,0.14862774],"category_scores_gemma":[0.034457102,0.00133586,0.002740436,0.004017394,0.0020044195,0.0025168255,0.0026952366,0.006597618,0.09070402],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019134452,0.000006602506,0.000025642847,0.00009089418,0.000009758328,0.000112752154,0.000011255516,0.00003863242,0.00010616244,0.0002288329,0.9961558,0.0031945284],"study_design_scores_gemma":[0.00003335985,0.000023989462,0.0012142343,0.0001713329,0.000050425755,0.0003861271,0.000052436186,0.00041917697,0.00081465166,0.0012027945,0.9955787,0.000052754145],"about_ca_topic_score_codex":0.028737107,"about_ca_topic_score_gemma":0.03825682,"teacher_disagreement_score":0.14862774,"about_ca_system_score_codex":0.0039641415,"about_ca_system_score_gemma":0.0031456128,"threshold_uncertainty_score":0.4972093},"labels":[],"label_agreement":null},{"id":"W2161172731","doi":"10.1145/1774088.1774420","title":"PROM-OOGLE","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Task (project management); Resource (disambiguation); Process (computing); Data science; Promoter; Biological database; Focus (optics); World Wide Web; Gene; Bioinformatics; Biology; Engineering; Genetics","score_opus":0.007801166297135559,"score_gpt":0.26093302127598494,"score_spread":0.2531318549788494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161172731","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056198006,0.001267925,0.27044705,0.0010574263,0.0004378824,0.0011203655,0.07597904,0.6126625,0.031408],"genre_scores_gemma":[0.06543105,0.0048119444,0.39956966,0.0031490745,0.0003821405,0.009227863,0.31207308,0.16313721,0.042217895],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979576,0.00048530527,0.0003045243,0.00045099092,0.0005635879,0.00023801434],"domain_scores_gemma":[0.99599385,0.0023519578,0.00034691757,0.0007224821,0.00031582476,0.00026893805],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.004819879,0.0028525505,0.002487353,0.0030874189,0.0011481785,0.00475467,0.004889418,0.0017799406,0.06321417],"category_scores_gemma":[0.009606576,0.0026480048,0.0026539315,0.003292989,0.0011562889,0.0045111096,0.0053896187,0.0037522146,0.038917635],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006128234,0.00052048225,0.005275341,0.0074642864,0.00072903885,0.0017848657,0.0018086352,0.0061287126,0.018973388,0.044934917,0.68732667,0.21892548],"study_design_scores_gemma":[0.0014123811,0.00025807956,0.0028337468,0.0007967125,0.00022159464,0.0011574004,0.00028087705,0.018445296,0.02041071,0.040940925,0.912977,0.00026517943],"about_ca_topic_score_codex":0.00166597,"about_ca_topic_score_gemma":0.0018797553,"teacher_disagreement_score":0.9367858,"about_ca_system_score_codex":0.00091525784,"about_ca_system_score_gemma":0.0028642195,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W2161580998","doi":"10.1186/1748-5908-5-78","title":"The trainees' perspective on developing an end-of-grant knowledge translation plan","year":2010,"lang":"en","type":"article","venue":"Implementation Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto; University of Manitoba; Queen's University; Toronto Metropolitan University; University of Calgary","funders":"Canadian Institutes of Health Research; Fondation pour la Recherche Médicale; European Observatory on Health Systems and Policies","keywords":"Knowledge translation; Medicine; Plan (archaeology); Craft; Perspective (graphical); Process (computing); Medical education; Developing country; Health informatics; Health services research; Knowledge management; Nursing; Public health; Computer science; Economic growth","score_opus":0.12887416057549947,"score_gpt":0.4666967740628391,"score_spread":0.33782261348733966,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2161580998","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014647739,0.0006556005,0.03299372,0.9192484,0.0018296686,0.0008700869,0.00007223496,0.00013644193,0.029546047],"genre_scores_gemma":[0.51270133,0.0029877122,0.18596946,0.26299712,0.0027181308,0.0038706295,0.00024288142,0.00023998288,0.028272772],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.73562145,0.2178653,0.0088431155,0.0035588918,0.015064483,0.019046687],"domain_scores_gemma":[0.70764446,0.17129712,0.012786238,0.009230281,0.038313814,0.060728088],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.24050306,0.00083053997,0.0007789024,0.0016011259,0.015093064,0.024444252,0.007364326,0.026957572,0.011087804],"category_scores_gemma":[0.2158502,0.0010382294,0.0017218802,0.00158447,0.017406711,0.014644445,0.021629192,0.034762528,0.0023257157],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004489867,0.0017438849,0.0069888537,0.0019532617,0.00014978982,0.0061958865,0.22143751,0.008464248,0.003050464,0.4301873,0.17267564,0.14670417],"study_design_scores_gemma":[0.00036406252,0.0010890699,0.0026014408,0.0030795478,0.00011017798,0.0018735928,0.24429274,0.008576449,0.0029513543,0.149527,0.58517104,0.00036363813],"about_ca_topic_score_codex":0.0084450925,"about_ca_topic_score_gemma":0.009370932,"teacher_disagreement_score":0.7594969,"about_ca_system_score_codex":0.015141106,"about_ca_system_score_gemma":0.12134784,"threshold_uncertainty_score":0.93659496},"labels":[],"label_agreement":null},{"id":"W2163269670","doi":"","title":"Combining relevance assignment with quality of the evidence to support guideline development.","year":2010,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Guideline; Computer science; Relevance (law); Recall; Quality (philosophy); Process (computing); Dissemination; Quality of evidence; Clinical Practice; Process management; Medicine; Psychology; Family medicine; Randomized controlled trial; Engineering; Pathology; Cognitive psychology","score_opus":0.05740096418627965,"score_gpt":0.3104970791983581,"score_spread":0.25309611501207846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163269670","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048720475,0.0073522977,0.8736985,0.012277253,0.0011777729,0.00982644,0.011055499,0.020376155,0.015515648],"genre_scores_gemma":[0.11766717,0.0008706445,0.87326366,0.0005302906,0.00035359213,0.0020741343,0.00419001,0.0003759124,0.0006746297],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9274464,0.02694159,0.021543527,0.0065699215,0.016785108,0.000713366],"domain_scores_gemma":[0.64799535,0.23901577,0.028000826,0.016332265,0.06595196,0.0027038413],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.083451934,0.001965755,0.0032550602,0.057230145,0.0025522364,0.008697145,0.0031063508,0.0022986105,0.0067672217],"category_scores_gemma":[0.40489414,0.0013194656,0.0033227233,0.022271328,0.001403153,0.008353661,0.007849098,0.0022117158,0.002475119],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072661537,0.0002610949,0.02629869,0.007960389,0.0010011521,0.0003128306,0.0020958579,0.0028221195,0.0046011247,0.0072032367,0.034729123,0.9119877],"study_design_scores_gemma":[0.0027178056,0.0011670355,0.08732561,0.0094779115,0.0094348965,0.0033708145,0.004999243,0.37273875,0.03814178,0.27882147,0.19034399,0.0014607116],"about_ca_topic_score_codex":0.004319603,"about_ca_topic_score_gemma":0.007607813,"teacher_disagreement_score":0.9165481,"about_ca_system_score_codex":0.0031645198,"about_ca_system_score_gemma":0.008658537,"threshold_uncertainty_score":0.44134128},"labels":[],"label_agreement":null},{"id":"W2163980860","doi":"10.1093/bioinformatics/bti493","title":"Discovering patterns to extract protein–protein interactions from the literature: Part II","year":2005,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":96,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Killam Trusts","keywords":"Merge (version control); Computer science; Recall; Protein–protein interaction; Generalization; Precision and recall; Protein expression; Artificial intelligence; Task (project management); Machine learning; Data mining; Computational biology; Information retrieval; Biology; Genetics; Mathematics","score_opus":0.015103357633382478,"score_gpt":0.26207536494041134,"score_spread":0.24697200730702887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2163980860","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15318106,0.0035386435,0.82659143,0.0022594861,0.00012559339,0.0011320775,0.005823164,0.0039294343,0.0034191243],"genre_scores_gemma":[0.16063589,0.0018094871,0.82468253,0.00029072815,0.00016040151,0.00087215164,0.009035481,0.00019808135,0.0023152656],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985436,0.00030309812,0.00029742194,0.0004024203,0.00037082753,0.00008267336],"domain_scores_gemma":[0.99719834,0.0014979693,0.00038027964,0.00042109683,0.00037260354,0.00012972135],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016087858,0.0010632095,0.00082314375,0.0069341017,0.0007705575,0.0019603923,0.0012879303,0.0009373401,0.0027673116],"category_scores_gemma":[0.00756898,0.00054148363,0.0011379122,0.0064156163,0.0009238071,0.002793149,0.0015196963,0.00083912956,0.0015565257],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023105262,0.0005541327,0.022902839,0.0015960323,0.00018224887,0.0007138824,0.0003979099,0.007193329,0.046551745,0.0049326466,0.008889494,0.90585476],"study_design_scores_gemma":[0.00039714776,0.00095811236,0.08505093,0.0006219078,0.00092431426,0.008364484,0.0016649974,0.5490792,0.19139782,0.087517135,0.07377168,0.00025230722],"about_ca_topic_score_codex":0.0015243152,"about_ca_topic_score_gemma":0.002191204,"teacher_disagreement_score":0.0069341017,"about_ca_system_score_codex":0.00049525616,"about_ca_system_score_gemma":0.0016714146,"threshold_uncertainty_score":0.009257615},"labels":[],"label_agreement":null},{"id":"W2164337735","doi":"10.1186/1472-6947-12-s1-s5","title":"Semantic text mining support for lignocellulose research","year":2012,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":53,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"Génome Québec; Genome Canada","keywords":"Computer science; Semantic Web; Lignocellulosic biomass; Biomass (ecology); World Wide Web; Biofuel; Data science; Biotechnology; Biology; Ecology","score_opus":0.09589379686879022,"score_gpt":0.4133700768724015,"score_spread":0.3174762800036113,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2164337735","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043669175,0.00465458,0.77485436,0.012987683,0.000660394,0.0021095688,0.07468902,0.06704799,0.01932718],"genre_scores_gemma":[0.10919141,0.0028708598,0.7823442,0.0014693909,0.00033741663,0.00089620653,0.09932418,0.0009137484,0.0026525296],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9967914,0.0008050516,0.0007694372,0.00064924697,0.0008781669,0.00010674048],"domain_scores_gemma":[0.9849764,0.008626575,0.0016476893,0.0017194147,0.002503538,0.0005262608],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053117466,0.0010463368,0.0007001631,0.008286864,0.0010724815,0.003086555,0.0018015775,0.0014056119,0.0058331625],"category_scores_gemma":[0.017637026,0.00032036853,0.0015331549,0.0061174156,0.00066375843,0.006342632,0.002907205,0.0012699837,0.0030052017],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070574036,0.0010121022,0.009502786,0.0071248026,0.00040892162,0.002891782,0.0015594529,0.016452795,0.03273205,0.06343757,0.098165624,0.76600647],"study_design_scores_gemma":[0.00028453153,0.00019363326,0.0060122753,0.00161942,0.00039328536,0.0021168587,0.0014632237,0.27362415,0.06398627,0.2537583,0.3964129,0.00013518205],"about_ca_topic_score_codex":0.0017668431,"about_ca_topic_score_gemma":0.0032975196,"teacher_disagreement_score":0.008286864,"about_ca_system_score_codex":0.0011659658,"about_ca_system_score_gemma":0.0028324488,"threshold_uncertainty_score":0.02809149},"labels":[],"label_agreement":null},{"id":"W2165013445","doi":"","title":"An Experiment on Using Temporal Ontologies to Reason about Localization and Transport of Fungal Proteins","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Executable; Hierarchy; Ontology; Computer science; Subcellular localization; Tree (set theory); Protein subcellular localization prediction; Computational biology; Biology; Programming language; Cell biology; Biochemistry; Mathematics; Gene; Cytoplasm","score_opus":0.024783587093264375,"score_gpt":0.3255114158556188,"score_spread":0.30072782876235443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165013445","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.51491815,0.00033869583,0.4358622,0.0028785742,0.00028900887,0.0012844582,0.006138309,0.02717623,0.011114421],"genre_scores_gemma":[0.53461474,0.00021197114,0.45558366,0.0004215672,0.00003759791,0.00055541395,0.0040133116,0.0011559796,0.0034057673],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99591774,0.0022798388,0.0004909428,0.0006011853,0.0005729585,0.00013733559],"domain_scores_gemma":[0.9599242,0.03453246,0.00046664145,0.0031096395,0.0013753583,0.0005916451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0093368655,0.0007123832,0.00049428793,0.0009004046,0.00083055184,0.0015817854,0.0012510193,0.0018840439,0.005444656],"category_scores_gemma":[0.025331989,0.0005223642,0.0012779197,0.0009064444,0.0010212123,0.0070367064,0.0019513562,0.0017150032,0.0011322291],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.020003658,0.0060930196,0.049753796,0.005654672,0.00096053723,0.004630569,0.018486962,0.060821418,0.19448149,0.07200776,0.03393237,0.5331738],"study_design_scores_gemma":[0.001976098,0.0026789424,0.0145981945,0.00035480454,0.00079102506,0.001838692,0.0037573008,0.62558866,0.21131745,0.028010422,0.108795375,0.00029312455],"about_ca_topic_score_codex":0.0074441014,"about_ca_topic_score_gemma":0.004205174,"teacher_disagreement_score":0.0093368655,"about_ca_system_score_codex":0.00090041413,"about_ca_system_score_gemma":0.0010841215,"threshold_uncertainty_score":0.049378633},"labels":[],"label_agreement":null},{"id":"W2165034468","doi":"10.1093/nar/gku956","title":"Xenbase, the Xenopus model organism database; new virtualized system, data types and genomes","year":2014,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":138,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; National Institute of Child Health and Human Development; National Institutes of Health","keywords":"Xenopus; Biology; Genome; Computational biology; Model organism; Organism; Database; Gene; Genetics; Computer science","score_opus":0.09007029162540417,"score_gpt":0.361542129770755,"score_spread":0.2714718381453508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165034468","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041211667,0.002565105,0.04787846,0.00059502595,0.00035894636,0.0002771147,0.824362,0.108713,0.0111292675],"genre_scores_gemma":[0.0075445473,0.0014092348,0.03363282,0.00024818562,0.00004824291,0.00036423773,0.94680434,0.007303254,0.0026450949],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992575,0.000108510154,0.00017290647,0.0001812205,0.00020840974,0.00007127245],"domain_scores_gemma":[0.9986467,0.0002552425,0.00022980425,0.00036631533,0.00019643552,0.00030564613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016230729,0.0018514955,0.001674821,0.0043288623,0.0007444935,0.003741007,0.0032923413,0.0012355449,0.027787242],"category_scores_gemma":[0.0039174426,0.001090606,0.0010051348,0.005038372,0.00044800888,0.00330654,0.0023970227,0.001555921,0.027305214],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026012356,0.000115957286,0.0037682531,0.0041469964,0.0003024057,0.0006968773,0.00034402355,0.0036505086,0.029186768,0.016267141,0.85287136,0.08604848],"study_design_scores_gemma":[0.00027119016,0.000070686176,0.0034726,0.000398227,0.00016862556,0.0005450219,0.00010664961,0.0037735105,0.01338167,0.0052126916,0.97248137,0.00011764151],"about_ca_topic_score_codex":0.005345626,"about_ca_topic_score_gemma":0.0044894777,"teacher_disagreement_score":0.027787242,"about_ca_system_score_codex":0.001186887,"about_ca_system_score_gemma":0.0023354648,"threshold_uncertainty_score":0.092957556},"labels":[],"label_agreement":null},{"id":"W2165057277","doi":"10.1186/s12859-015-0488-1","title":"OTO: Ontology Term Organizer","year":2015,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"National Science Foundation","keywords":"Term (time); Ontology; Computer science; Computational biology; Information retrieval; DNA microarray; Gene ontology; Biology; Data science; Bioinformatics; World Wide Web; Genetics; Gene; Physics","score_opus":0.04411125647240074,"score_gpt":0.2846898439338195,"score_spread":0.24057858746141877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165057277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014829662,0.0014850679,0.5357282,0.0016956708,0.00077865104,0.0032160187,0.23069547,0.18002288,0.031548362],"genre_scores_gemma":[0.030814456,0.0009939584,0.5297168,0.0006156104,0.00023377,0.0036529456,0.40114912,0.02406706,0.008756206],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99690443,0.0003740743,0.00073132775,0.0007269803,0.0010064445,0.00025669788],"domain_scores_gemma":[0.9928068,0.0023629174,0.0013483324,0.0014737444,0.0013646653,0.0006434326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035898893,0.002112391,0.0010617848,0.010953894,0.0018338602,0.0041427976,0.0024375175,0.0010442641,0.018666627],"category_scores_gemma":[0.017008262,0.0011390537,0.0023963763,0.010506304,0.001234324,0.0070508546,0.0062422855,0.0023234366,0.01448491],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076213863,0.00021081467,0.011941557,0.0049069338,0.00026602278,0.0008993078,0.004671645,0.004025831,0.015551549,0.045128174,0.57019234,0.3414437],"study_design_scores_gemma":[0.00014506067,0.0000950252,0.007420312,0.00074107863,0.00012781053,0.0008453059,0.0011246693,0.0134976925,0.00942992,0.046197705,0.920182,0.0001935071],"about_ca_topic_score_codex":0.009752534,"about_ca_topic_score_gemma":0.010122339,"teacher_disagreement_score":0.018666627,"about_ca_system_score_codex":0.0018180998,"about_ca_system_score_gemma":0.005983253,"threshold_uncertainty_score":0.062446058},"labels":[],"label_agreement":null},{"id":"W2165336562","doi":"10.1186/1756-0381-5-14","title":"Peer2ref: a peer-reviewer finding web tool that uses author disambiguation","year":2012,"lang":"en","type":"article","venue":"BioData Mining","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"","keywords":"Computer science; MEDLINE; Information retrieval; Web of science; World Wide Web; Selection (genetic algorithm); Data science; Artificial intelligence","score_opus":0.11928096650161898,"score_gpt":0.35438555593076476,"score_spread":0.23510458942914578,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165336562","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008359505,0.0025768478,0.6225388,0.0027538764,0.0024369925,0.003578244,0.0436114,0.29685122,0.017293163],"genre_scores_gemma":[0.02903231,0.001336534,0.89727974,0.0007546177,0.001304763,0.0023263409,0.035858165,0.017180827,0.014926655],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9847112,0.0042964355,0.0031951822,0.0024624118,0.004904316,0.0004304615],"domain_scores_gemma":[0.88454646,0.05935724,0.012727747,0.012227881,0.026677586,0.0044630254],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.029585812,0.0026977742,0.0026570407,0.033184923,0.003908306,0.007577324,0.0039202473,0.0032650104,0.056233842],"category_scores_gemma":[0.0927793,0.0018520666,0.0014145694,0.01518847,0.0012672003,0.011008887,0.0065397206,0.0016665538,0.045649014],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090048235,0.0002698904,0.0069876355,0.006407685,0.00039216093,0.0009537725,0.0020089068,0.001044661,0.012126841,0.010011188,0.38833225,0.5705645],"study_design_scores_gemma":[0.0005281137,0.00026436764,0.0070669292,0.0013379657,0.00029726262,0.0018224155,0.0010998503,0.024359163,0.029721823,0.022562288,0.91035837,0.00058148656],"about_ca_topic_score_codex":0.0015823059,"about_ca_topic_score_gemma":0.0028074342,"teacher_disagreement_score":0.97041416,"about_ca_system_score_codex":0.0010734668,"about_ca_system_score_gemma":0.008008891,"threshold_uncertainty_score":0.1881209},"labels":[],"label_agreement":null},{"id":"W2165428180","doi":"10.1016/j.jbi.2008.03.004","title":"Bio2RDF: Towards a mashup to build bioinformatics knowledge systems","year":2008,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":798,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University; Université Laval; Centre hospitalier de l'Université Laval","funders":"Canadian Institutes of Health Research; Génome Québec; Genome Canada","keywords":"Computer science; Mashup; Bioinformatics; World Wide Web; Data science; The Internet; Biology; Web 2.0","score_opus":0.025953939340733118,"score_gpt":0.29290583438515777,"score_spread":0.2669518950444246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2165428180","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045517003,0.0003335802,0.9167484,0.0006115412,0.00013719956,0.00027543984,0.00310948,0.072511986,0.0017206358],"genre_scores_gemma":[0.020304793,0.00036803156,0.959418,0.00047791572,0.00007141138,0.00032321518,0.011523637,0.004801865,0.002711157],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99827194,0.0003309164,0.00020028758,0.00047560703,0.0005947379,0.00012641537],"domain_scores_gemma":[0.99615806,0.0017580456,0.00020278002,0.0011258537,0.00040224148,0.00035306896],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003549929,0.001456194,0.0014941506,0.0028479928,0.0013549248,0.003780636,0.0036020144,0.0022510192,0.005852779],"category_scores_gemma":[0.008193845,0.0012540634,0.0023305402,0.0021138126,0.0011242556,0.0068570655,0.0041488414,0.0031316024,0.0031419504],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092987827,0.00078430324,0.0038472223,0.0023722844,0.00084308354,0.0019840328,0.0034095084,0.023988986,0.06508697,0.061521787,0.12675367,0.7084783],"study_design_scores_gemma":[0.0004127058,0.0002442798,0.002251071,0.00067717215,0.0005421129,0.0016379771,0.00069356145,0.3282677,0.090774015,0.15887222,0.4152949,0.00033233964],"about_ca_topic_score_codex":0.005211261,"about_ca_topic_score_gemma":0.0060785445,"teacher_disagreement_score":0.005852779,"about_ca_system_score_codex":0.0010793873,"about_ca_system_score_gemma":0.0017164438,"threshold_uncertainty_score":0.01957953},"labels":[],"label_agreement":null},{"id":"W2166240318","doi":"10.1093/database/bas056","title":"An overview of the BioCreative 2012 Workshop Track III: interactive text mining task","year":2013,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute for Research in Immunology and Cancer","funders":"U.S. National Library of Medicine; National Human Genome Research Institute","keywords":"Computer science; Data curation; Annotation; Task (project management); Usability; Information retrieval; Data science; Biomedical text mining; Set (abstract data type); World Wide Web; Text mining; Natural language processing; Artificial intelligence; Human–computer interaction; Engineering","score_opus":0.0459006993030045,"score_gpt":0.3375506222957032,"score_spread":0.2916499229926987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166240318","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06780125,0.014673149,0.44403687,0.011951915,0.00519114,0.016415225,0.18228161,0.2028161,0.054832704],"genre_scores_gemma":[0.056534875,0.0030029763,0.49981785,0.00364145,0.0006989037,0.016897302,0.34972548,0.014385231,0.055296],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98964816,0.0026739116,0.0012422174,0.0030248591,0.002667575,0.00074338075],"domain_scores_gemma":[0.9842606,0.0060615833,0.00056260946,0.0026217138,0.005053025,0.0014404637],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014722865,0.0039619636,0.002962233,0.0045310957,0.0047710366,0.005669924,0.0065414086,0.0037722278,0.028855162],"category_scores_gemma":[0.022547869,0.0016435088,0.0032604374,0.0056257234,0.00094138685,0.006200093,0.0062139765,0.0036581913,0.03287193],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012504102,0.0011072248,0.0032956034,0.0041033425,0.00019996594,0.0004895073,0.0029159805,0.0019793215,0.026182737,0.0017668818,0.65146387,0.30524513],"study_design_scores_gemma":[0.0008209833,0.0011639323,0.014740254,0.0008729506,0.00028852595,0.0009937978,0.0019251445,0.031373702,0.055210017,0.0068014152,0.8853098,0.00049954135],"about_ca_topic_score_codex":0.0151302805,"about_ca_topic_score_gemma":0.018797882,"teacher_disagreement_score":0.028855162,"about_ca_system_score_codex":0.0032270534,"about_ca_system_score_gemma":0.0059923353,"threshold_uncertainty_score":0.0965302},"labels":[],"label_agreement":null},{"id":"W2166516661","doi":"10.1093/database/bas017","title":"How to link ontologies and protein-protein interactions to literature: text-mining approaches and the BioCreative experience","year":2012,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":34,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute for Research in Immunology and Cancer","funders":"Biotechnology and Biological Sciences Research Council; National Center for Research Resources; Directorate for Biological Sciences; National Institutes of Health; Office of Research Infrastructure Programs, National Institutes of Health; Canadian Institutes of Health Research; European Commission; Wellcome Trust","keywords":"Computer science; Ontology; Information retrieval; Pipeline (software); Context (archaeology); Data curation; Information extraction; Consistency (knowledge bases); Workflow; Open Biomedical Ontologies; Biomedical text mining; Controlled vocabulary; Process (computing); Annotation; Vocabulary; Natural language processing; Upper ontology; Data science; Suggested Upper Merged Ontology; Artificial intelligence; Text mining; Semantic Web; Database","score_opus":0.045486472278633866,"score_gpt":0.294178056247657,"score_spread":0.24869158396902316,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166516661","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026268872,0.06489862,0.78902006,0.06287145,0.0015492897,0.0010708706,0.0070309797,0.0031711725,0.04411873],"genre_scores_gemma":[0.07798959,0.033353418,0.8700514,0.004443796,0.0011903633,0.0007403518,0.004905905,0.00081325544,0.0065120743],"study_design_codex":"design_other","study_design_gemma":"design_other","domain_scores_codex":[0.98950773,0.0049420986,0.0018088757,0.0010784836,0.002459156,0.00020361955],"domain_scores_gemma":[0.94926924,0.037883665,0.0034090162,0.0029713474,0.005687299,0.0007794268],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017994752,0.0010653986,0.0013331637,0.047663927,0.0023742826,0.012243057,0.0025593732,0.0022079255,0.005392216],"category_scores_gemma":[0.055367894,0.00077948754,0.0014816448,0.027758587,0.004947551,0.022854716,0.004310654,0.0022886838,0.002775411],"study_design_candidate":"design_other","study_design_consensus":"design_other","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007429823,0.00016516577,0.0044320063,0.006444501,0.00025647428,0.0017093465,0.0153052155,0.0017301318,0.00718023,0.16937564,0.042595156,0.75073177],"study_design_scores_gemma":[0.00003540733,0.000042889515,0.0045387023,0.0064894166,0.00017519502,0.0019417213,0.012165451,0.010038312,0.0045448914,0.39960223,0.5602973,0.00012850098],"about_ca_topic_score_codex":0.0034099533,"about_ca_topic_score_gemma":0.005008338,"teacher_disagreement_score":0.047663927,"about_ca_system_score_codex":0.0023880065,"about_ca_system_score_gemma":0.0030158756,"threshold_uncertainty_score":0.095166445},"labels":[],"label_agreement":null},{"id":"W2166588833","doi":"10.1016/j.jalz.2013.12.013","title":"International Alzheimer's Disease Research Portfolio (IADRP) aims to capture global Alzheimer's disease research funding","year":2014,"lang":"en","type":"article","venue":"Alzheimer s & Dementia","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Leverage (statistics); Portfolio; Disease; Business; Alzheimer's disease; Public relations; Political science; Medicine; Finance; Computer science","score_opus":0.09018658180639445,"score_gpt":0.39592683640528936,"score_spread":0.3057402545988949,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166588833","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060975097,0.012912895,0.16006655,0.03674849,0.0019262879,0.0067742504,0.31003138,0.0049366015,0.40562835],"genre_scores_gemma":[0.14871363,0.013431789,0.4337505,0.0037646494,0.0007714075,0.0071196253,0.3673043,0.0012790893,0.02386506],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9856838,0.0047709546,0.0025259706,0.0015051805,0.0043883557,0.0011257597],"domain_scores_gemma":[0.95482606,0.00933078,0.010341847,0.0054480713,0.014396193,0.0056570484],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.03166307,0.0015957223,0.0014507179,0.024591815,0.0023659258,0.009264993,0.0024280883,0.0018777604,0.009159114],"category_scores_gemma":[0.05437183,0.000684703,0.0016345629,0.038422268,0.0010217323,0.010379936,0.011139647,0.0022264484,0.006714216],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040566546,0.0004341464,0.06301737,0.0027632646,0.00038068386,0.00024145815,0.0021140606,0.0026287886,0.0014304735,0.11170241,0.29380134,0.5210804],"study_design_scores_gemma":[0.00009203091,0.00015919571,0.053634223,0.002025101,0.00025781407,0.0003938464,0.0020677012,0.0033483456,0.001585106,0.045886513,0.8904687,0.00008134761],"about_ca_topic_score_codex":0.015250577,"about_ca_topic_score_gemma":0.0103526935,"teacher_disagreement_score":0.9754082,"about_ca_system_score_codex":0.004755674,"about_ca_system_score_gemma":0.02455822,"threshold_uncertainty_score":0.16745234},"labels":[],"label_agreement":null},{"id":"W2166596803","doi":"10.1038/sj.embor.embor833","title":"The way we write","year":2003,"lang":"en","type":"article","venue":"EMBO Reports","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":29,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Children’s Health Research Institute","funders":"","keywords":"Computer science; Computational biology; Business; Biology","score_opus":0.013013483365418396,"score_gpt":0.2606692196357895,"score_spread":0.24765573627037107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166596803","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034881646,0.009426213,0.040987503,0.4091729,0.06959624,0.00015004289,0.0018386876,0.0016526006,0.46368763],"genre_scores_gemma":[0.094701745,0.00931938,0.03867692,0.10130307,0.01717232,0.00047894777,0.0024272935,0.002819095,0.73310125],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9880165,0.005975527,0.0009431607,0.0013825895,0.002977475,0.00070466317],"domain_scores_gemma":[0.97679615,0.005545729,0.0024318458,0.004765284,0.007830936,0.0026300978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0096187275,0.0006495744,0.0005238056,0.0021658835,0.003972147,0.018014746,0.0015408209,0.0034222715,0.04219594],"category_scores_gemma":[0.035320465,0.00051415985,0.0007194918,0.0019932352,0.010456048,0.016624136,0.0052472176,0.009471844,0.039755337],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038156297,0.000019379502,0.0005034829,0.0001389011,0.000026971746,0.00014458685,0.00741483,0.000044415996,0.0005882854,0.41534948,0.53476596,0.040965542],"study_design_scores_gemma":[0.0000041832673,0.000004259409,0.00012081165,0.00007118642,0.0000060244674,0.00007459439,0.0017073697,0.000025674786,0.00016445368,0.03036363,0.96744776,0.00001012668],"about_ca_topic_score_codex":0.0042374586,"about_ca_topic_score_gemma":0.004446186,"teacher_disagreement_score":0.04219594,"about_ca_system_score_codex":0.0022266812,"about_ca_system_score_gemma":0.00665009,"threshold_uncertainty_score":0.14115947},"labels":[],"label_agreement":null},{"id":"W2166648275","doi":"10.2196/medinform.4959","title":"NHash: Randomized N-Gram Hashing for Distributed Generation of Validatable Unique Study Identifiers in Multicenter Research","year":2015,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; National Institute of Neurological Disorders and Stroke; National Heart, Lung, and Blood Institute","keywords":"Identifier; Computer science; Hash function; Unique identifier; n-gram; Cryptography; Universal hashing; Theoretical computer science; Hash table; Data mining; Computer security; Computer network; Artificial intelligence; Double hashing","score_opus":0.1453992460823107,"score_gpt":0.4347469399101895,"score_spread":0.2893476938278788,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2166648275","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019920371,0.00030059484,0.9687275,0.00038390936,0.00026529247,0.0014317455,0.0005771688,0.00669698,0.0016964257],"genre_scores_gemma":[0.20499241,0.00014351548,0.78837657,0.00030604843,0.00012933434,0.0020452996,0.00087457185,0.00043201525,0.002700166],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98863226,0.0063299397,0.00090191315,0.0013618742,0.0023068,0.0004671845],"domain_scores_gemma":[0.97522527,0.009261378,0.003188506,0.009191133,0.0021799407,0.0009537417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011880726,0.0008285961,0.0009135398,0.0012286905,0.0012656853,0.0013456316,0.0025368822,0.0011871753,0.008436088],"category_scores_gemma":[0.038609684,0.00066724134,0.0008139419,0.0013135225,0.0018408089,0.0026986585,0.005504002,0.0011220789,0.004071686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006663168,0.0005170885,0.009282849,0.0013139652,0.00029261733,0.0013498655,0.0026109482,0.03153176,0.06257721,0.07198426,0.026784018,0.7850922],"study_design_scores_gemma":[0.002971551,0.0041411542,0.006202728,0.00042294367,0.0002548321,0.0027724975,0.0013451457,0.56202656,0.16213842,0.15589856,0.10133821,0.00048743724],"about_ca_topic_score_codex":0.0006604015,"about_ca_topic_score_gemma":0.00088619156,"teacher_disagreement_score":0.011880726,"about_ca_system_score_codex":0.0010545629,"about_ca_system_score_gemma":0.0030375305,"threshold_uncertainty_score":0.06283206},"labels":[],"label_agreement":null},{"id":"W2169250641","doi":"10.1186/2041-1480-2-s2-s10","title":"Integration and publication of heterogeneous text-mined relationships on the Semantic Web","year":2011,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":41,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"U.S. National Library of Medicine; National Human Genome Research Institute; National Institutes of Health; National Science Foundation","keywords":"Ontology; Computer science; WordNet; Semantics (computer science); SPARQL; Semantic Web; Information retrieval; Syntax; Upper ontology; RDF; Natural language processing; World Wide Web; Programming language","score_opus":0.048755736780086545,"score_gpt":0.2608897797600274,"score_spread":0.21213404297994087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169250641","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.122261934,0.0041331127,0.78784686,0.005185952,0.0006823657,0.000799432,0.024987953,0.015330234,0.038772076],"genre_scores_gemma":[0.2787653,0.0037245248,0.664536,0.0008937887,0.0002851796,0.00041378295,0.045235585,0.0019867257,0.004159154],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99335915,0.0019599688,0.0010848147,0.00092705217,0.002427199,0.00024180475],"domain_scores_gemma":[0.9802215,0.00917672,0.0023864217,0.004485727,0.0032957315,0.0004339472],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008371523,0.0006845287,0.0008648699,0.01383593,0.0011987197,0.0048947777,0.0012666049,0.0013365569,0.0026376634],"category_scores_gemma":[0.022547916,0.000475106,0.0015405568,0.016041253,0.0011021486,0.009258443,0.003359117,0.0013643386,0.0015486844],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000610748,0.00065509137,0.021725968,0.003606327,0.0007128155,0.0057493076,0.004564885,0.018173905,0.045659065,0.17609046,0.047207106,0.67524433],"study_design_scores_gemma":[0.00010598964,0.00019735526,0.016619723,0.0015949588,0.0007155192,0.0029046317,0.0026627695,0.12607369,0.09147865,0.22941288,0.5279963,0.00023758232],"about_ca_topic_score_codex":0.0016871627,"about_ca_topic_score_gemma":0.001890978,"teacher_disagreement_score":0.01383593,"about_ca_system_score_codex":0.0011913797,"about_ca_system_score_gemma":0.0025692815,"threshold_uncertainty_score":0.044273317},"labels":[],"label_agreement":null},{"id":"W2169425537","doi":"10.1093/database/bau067","title":"Assisting manual literature curation for protein-protein interactions using BioQRator","year":2014,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Institute for Research in Immunology and Cancer","funders":"U.S. National Library of Medicine; National Institutes of Health; Ministry of Education, Science and Technology; National Research Foundation of Korea; National Research Foundation","keywords":"Annotation; Computer science; Data curation; Task (project management); Information retrieval; World Wide Web; Artificial intelligence","score_opus":0.030990985643682793,"score_gpt":0.3347915863824545,"score_spread":0.3038006007387717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2169425537","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034790702,0.019985508,0.58021075,0.0080769975,0.0023549155,0.009119655,0.12792274,0.19128886,0.026249906],"genre_scores_gemma":[0.025685517,0.004225966,0.88085234,0.0012767988,0.00037944678,0.0037323246,0.070476376,0.005985632,0.0073855687],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.97462374,0.011908915,0.0060562347,0.0029674436,0.0039109187,0.00053278625],"domain_scores_gemma":[0.80969924,0.11087504,0.014512337,0.023002647,0.039082434,0.0028283452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.046210412,0.0022792318,0.0036799735,0.036880523,0.0027725592,0.006884377,0.0033423572,0.001456599,0.035865944],"category_scores_gemma":[0.122582205,0.0012658915,0.0026639237,0.01974406,0.0010385811,0.0060399566,0.008058565,0.0019429853,0.034864534],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091344194,0.0003433437,0.007555298,0.032367304,0.0007644474,0.0012741798,0.0051932298,0.00083119335,0.04489646,0.0040295594,0.2701219,0.6317097],"study_design_scores_gemma":[0.00040732318,0.0003256063,0.014505694,0.005728619,0.0010317122,0.0023271898,0.002732445,0.009818451,0.03924037,0.0077654556,0.91564876,0.0004683459],"about_ca_topic_score_codex":0.0036915971,"about_ca_topic_score_gemma":0.009801192,"teacher_disagreement_score":0.046210412,"about_ca_system_score_codex":0.001607876,"about_ca_system_score_gemma":0.0110325655,"threshold_uncertainty_score":0.24438691},"labels":[],"label_agreement":null},{"id":"W2170282111","doi":"10.1093/nar/gkt1026","title":"The Human Phenotype Ontology project: linking molecular biology and disease through phenotype data","year":2013,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":837,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children; University of Toronto","funders":"Basic Energy Sciences; National Institute of Mental Health; National Institutes of Health; National Human Genome Research Institute; Bundesministerium für Bildung und Forschung; Office of Science; University College London; British Heart Foundation; National Institute for Health and Care Research; Deutsche Forschungsgemeinschaft; U.S. Department of Energy","keywords":"Unified Medical Language System; Annotation; Ontology; DECIPHER; Controlled vocabulary; Documentation; Computer science; Biology; Phenotype; Set (abstract data type); Interoperability; UniProt; Computational biology; Function (biology); Information retrieval; Bioinformatics; World Wide Web; Genetics","score_opus":0.09187093386330562,"score_gpt":0.41764594534052984,"score_spread":0.3257750114772242,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2170282111","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01379814,0.003377949,0.5504738,0.005216989,0.0006557177,0.0020253945,0.38289502,0.018063158,0.023493886],"genre_scores_gemma":[0.03480443,0.0056773913,0.4497757,0.0017214365,0.00021978978,0.0032328044,0.4977994,0.0027730775,0.0039959396],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971335,0.00089792424,0.00064359844,0.00045408518,0.0007195996,0.00015127451],"domain_scores_gemma":[0.99439764,0.0024877773,0.0007891236,0.0011742936,0.0006366065,0.00051457365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060293227,0.0015037168,0.0009770232,0.0079364395,0.0010734557,0.0024323494,0.0018642932,0.0012456706,0.0070829685],"category_scores_gemma":[0.010623163,0.0007680214,0.0016336326,0.009203642,0.0011069445,0.004292832,0.0051596626,0.0019689521,0.002835512],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065052806,0.00039083775,0.016182853,0.007679265,0.0008651827,0.0018245345,0.0036923862,0.006246996,0.02017593,0.1834766,0.41425338,0.34456155],"study_design_scores_gemma":[0.0002588813,0.000083950974,0.017586278,0.0017377387,0.0003282403,0.0013851918,0.00083277764,0.008713572,0.0057544955,0.08019108,0.8829812,0.00014668352],"about_ca_topic_score_codex":0.012274703,"about_ca_topic_score_gemma":0.0076902667,"teacher_disagreement_score":0.012274703,"about_ca_system_score_codex":0.0015181552,"about_ca_system_score_gemma":0.0081569,"threshold_uncertainty_score":0.03188646},"labels":[],"label_agreement":null},{"id":"W2171188883","doi":"10.1109/cbms.2007.79","title":"Ontology Engineering to Model Clinical Pathways: Towards the Computerization and Execution of Clinical Pathways","year":2007,"lang":"en","type":"article","venue":"Proceedings - IEEE Symposium on Computer-Based Medical Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Ontology; Computer science; Clinical pathway; Abstraction; Software engineering; Process (computing); Knowledge management; Data science; Medicine; Programming language","score_opus":0.06430275580445946,"score_gpt":0.32361946449489265,"score_spread":0.2593167086904332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171188883","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006317702,0.00012350958,0.98876125,0.0009406003,0.000029483066,0.0003844952,0.0008255958,0.0007960485,0.001821282],"genre_scores_gemma":[0.047679145,0.00027532826,0.94917464,0.00012964907,0.000010813991,0.00044163785,0.0016876626,0.000088118286,0.000512973],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970017,0.0013458019,0.00046657253,0.00043521702,0.0006294486,0.00012121473],"domain_scores_gemma":[0.9928924,0.004033148,0.0006790939,0.0010902255,0.0011123403,0.0001928842],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045720907,0.00070536963,0.00055843283,0.003359125,0.0013011718,0.0041266815,0.0017167415,0.0011643465,0.0016708211],"category_scores_gemma":[0.017773813,0.00062868063,0.0021786606,0.0034736125,0.0016266489,0.0041686664,0.002514209,0.002069751,0.0005177327],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018312478,0.0003019043,0.010817926,0.0011157194,0.00033527316,0.0008121104,0.0064088833,0.29747993,0.0045993268,0.38701123,0.009122703,0.28181177],"study_design_scores_gemma":[0.000078446625,0.00006248205,0.0013006974,0.0005129757,0.00017704134,0.0003201582,0.0014172941,0.61040723,0.0048047,0.3192078,0.061635766,0.00007539361],"about_ca_topic_score_codex":0.026705498,"about_ca_topic_score_gemma":0.028036177,"teacher_disagreement_score":0.026705498,"about_ca_system_score_codex":0.002872397,"about_ca_system_score_gemma":0.007979758,"threshold_uncertainty_score":0.05310011},"labels":[],"label_agreement":null},{"id":"W2171970585","doi":"10.1109/dial.2006.45","title":"Use of Figures in Literature Mining for Biomedical Digital Libraries","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Annotation; Computer science; Information retrieval; Biomedical text mining; Classifier (UML); Triage; Task (project management); Naive Bayes classifier; Metadata; Controlled vocabulary; Process (computing); Natural language processing; Text mining; Artificial intelligence; World Wide Web","score_opus":0.018578927195849724,"score_gpt":0.25023132423815386,"score_spread":0.23165239704230414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171970585","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09124405,0.015664142,0.67471176,0.0052702148,0.0012806995,0.005001218,0.09274684,0.0748814,0.039199702],"genre_scores_gemma":[0.09995199,0.0027770882,0.84754723,0.00030877878,0.00025199555,0.0010201397,0.04351988,0.0011315494,0.003491334],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9913571,0.0014573657,0.0016462847,0.0020893137,0.0031132968,0.00033669427],"domain_scores_gemma":[0.9656037,0.015624237,0.0045631244,0.005526504,0.007436465,0.0012458576],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.00780237,0.0013001845,0.0016134812,0.06488569,0.0023826791,0.007396776,0.002617452,0.0013229487,0.011354148],"category_scores_gemma":[0.05911453,0.0007929505,0.002134764,0.038951136,0.0014057102,0.007906079,0.0028350991,0.0014700742,0.008261473],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003358618,0.00024284438,0.018949732,0.0030975805,0.00027473248,0.0011192425,0.0018180151,0.00574662,0.009307961,0.014256,0.06840918,0.8764422],"study_design_scores_gemma":[0.00024970106,0.0006272433,0.043286078,0.0031024173,0.00100658,0.006579927,0.005502138,0.17368865,0.054848116,0.12314035,0.5874763,0.0004925849],"about_ca_topic_score_codex":0.004559508,"about_ca_topic_score_gemma":0.00664947,"teacher_disagreement_score":0.9351143,"about_ca_system_score_codex":0.0021308651,"about_ca_system_score_gemma":0.0040780683,"threshold_uncertainty_score":0.041263342},"labels":[],"label_agreement":null},{"id":"W2172304783","doi":"10.3233/ao-2011-0086","title":"Overcoming the ontology enrichment bottleneck with Quick Term Templates","year":2011,"lang":"en","type":"article","venue":"Applied Ontology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Cancer Agency","funders":"Biotechnology and Biological Sciences Research Council","keywords":"Computer science; Bottleneck; Term (time); Template; Ontology; Information retrieval; Programming language; Embedded system","score_opus":0.022578760276110207,"score_gpt":0.24094874650184145,"score_spread":0.21836998622573123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2172304783","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012569165,0.0005326788,0.9541321,0.0012651314,0.00029135824,0.00062698225,0.002636324,0.023626003,0.0043203887],"genre_scores_gemma":[0.03727899,0.0005218042,0.94735855,0.0005600344,0.00009539931,0.00036293262,0.0054735015,0.003647277,0.004701528],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9940568,0.0017401231,0.0012997211,0.0007866703,0.0018570555,0.00025964784],"domain_scores_gemma":[0.9559763,0.024497611,0.0022877858,0.009894495,0.006679596,0.00066417566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008651262,0.0014713381,0.0014681107,0.004422141,0.0011334037,0.004955617,0.0031689496,0.001617663,0.007557817],"category_scores_gemma":[0.041717265,0.0015780703,0.002185519,0.0047180224,0.00080671255,0.0074653407,0.0040946174,0.003912065,0.006871684],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063063897,0.00050073967,0.0056287036,0.0020728242,0.00026242,0.0012585965,0.002109655,0.005837309,0.05884307,0.042995546,0.05664329,0.8232174],"study_design_scores_gemma":[0.00025233236,0.00030702006,0.0034329293,0.0008648658,0.00063762785,0.003425176,0.0014053449,0.1632337,0.19734281,0.11123935,0.517574,0.00028488404],"about_ca_topic_score_codex":0.0031426132,"about_ca_topic_score_gemma":0.0055753365,"teacher_disagreement_score":0.008651262,"about_ca_system_score_codex":0.0010892759,"about_ca_system_score_gemma":0.0047037583,"threshold_uncertainty_score":0.045752764},"labels":[],"label_agreement":null},{"id":"W2175612905","doi":"10.1007/978-3-642-03262-2_5","title":"A Conceptual Framework for Ontology Based Automating and Merging of Clinical Pathways of Comorbidities","year":2009,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Ontology; Software engineering; Information retrieval; Data science; Epistemology","score_opus":0.0533534947781595,"score_gpt":0.3349154529542947,"score_spread":0.2815619581761352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2175612905","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018142856,0.00042153764,0.9850048,0.00076292397,0.00009983316,0.00043493215,0.0028257698,0.006603892,0.0020320157],"genre_scores_gemma":[0.01200369,0.0003257786,0.9815252,0.00022981162,0.00002875703,0.00025715207,0.0046974365,0.00025512604,0.0006769462],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9957742,0.0007253034,0.0010990173,0.0008753044,0.0012526539,0.00027353654],"domain_scores_gemma":[0.9942918,0.0020229158,0.0005582084,0.0008782366,0.0017743455,0.00047451287],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008319531,0.0013166468,0.0015478728,0.007733207,0.002400833,0.0071748267,0.004282881,0.0022189228,0.004557128],"category_scores_gemma":[0.011987443,0.0011450715,0.004498303,0.007186854,0.0014743998,0.0065604323,0.0053394255,0.0030200877,0.0022232123],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028739628,0.00042329406,0.0065876725,0.0021304626,0.00074933737,0.0018729736,0.0043395604,0.02655961,0.010733926,0.37600636,0.048538033,0.5217714],"study_design_scores_gemma":[0.00015400282,0.00014809429,0.0035033284,0.0017757501,0.0011514358,0.0021133274,0.0016285442,0.21929757,0.015533512,0.35341287,0.40100127,0.00028032533],"about_ca_topic_score_codex":0.028407643,"about_ca_topic_score_gemma":0.037084073,"teacher_disagreement_score":0.028407643,"about_ca_system_score_codex":0.0026205685,"about_ca_system_score_gemma":0.009518299,"threshold_uncertainty_score":0.05648458},"labels":[],"label_agreement":null},{"id":"W2179813712","doi":"10.3233/ao-2010-0081","title":"Open Biomedical Ontologies applied to prostate cancer","year":2011,"lang":"en","type":"article","venue":"Applied Ontology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"London Health Sciences Centre; Western University","funders":"","keywords":"Computer science; Prostate cancer; Data science; Information retrieval; Cancer; Medicine; Internal medicine","score_opus":0.045742484265470534,"score_gpt":0.3102968715088134,"score_spread":0.2645543872433429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2179813712","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13585334,0.015399984,0.7652349,0.008537282,0.00084375823,0.0011983409,0.0238929,0.004795165,0.044244446],"genre_scores_gemma":[0.5650745,0.012034652,0.37800017,0.0019149639,0.00030646595,0.0006930194,0.037008405,0.00080519915,0.004162741],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9919511,0.0029096769,0.0009823199,0.0009641998,0.0027850706,0.00040774373],"domain_scores_gemma":[0.98458767,0.00931119,0.0012271884,0.00212874,0.0023386795,0.00040648432],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004735996,0.0006897673,0.0006535343,0.010324926,0.0019475141,0.0035688828,0.0009374923,0.00087391504,0.00241106],"category_scores_gemma":[0.03026236,0.0003298806,0.0014905994,0.011853222,0.0012403664,0.0044990545,0.0042388593,0.0013908128,0.0006745472],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047816618,0.0004447649,0.039276842,0.0048088375,0.0007264775,0.002037365,0.005952268,0.03030632,0.017963663,0.21392737,0.023615094,0.6604628],"study_design_scores_gemma":[0.00010251576,0.00015740813,0.041716743,0.0023289542,0.0007797586,0.0025874926,0.0037538284,0.092932925,0.020676287,0.29602847,0.5387457,0.00018995558],"about_ca_topic_score_codex":0.016333522,"about_ca_topic_score_gemma":0.01744833,"teacher_disagreement_score":0.016333522,"about_ca_system_score_codex":0.0031951016,"about_ca_system_score_gemma":0.0054227468,"threshold_uncertainty_score":0.032476902},"labels":[],"label_agreement":null},{"id":"W2180878254","doi":"10.5167/uzh-64476","title":"Proceedings of the 5th International Symposium on Semantic Mining in Biomedicine (SMBM 2012)","year":2012,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Biotechnology and Biological Sciences Research Council; Directorate for Biological Sciences; New Brunswick Innovation Foundation","keywords":"Computer science; Ambiguity; Identification (biology); Set (abstract data type); Software; Task (project management); Matching (statistics); Information retrieval; Database; Baseline (sea); Programming language; Engineering","score_opus":0.06056999335213596,"score_gpt":0.30987421315309993,"score_spread":0.24930421980096396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2180878254","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02188482,0.15261376,0.5827121,0.0519648,0.06290647,0.0011234634,0.0072551644,0.013529991,0.10600947],"genre_scores_gemma":[0.065962516,0.07788092,0.49119893,0.010280092,0.016470782,0.0010083831,0.031040609,0.0040751416,0.3020826],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99679154,0.0012706411,0.00030063913,0.000489645,0.0009175165,0.00023010718],"domain_scores_gemma":[0.9943725,0.0015027359,0.00021635671,0.0012474133,0.0015440861,0.0011168822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00888466,0.0014418454,0.0018710437,0.0032377478,0.0011071786,0.00608552,0.0022826814,0.0030608212,0.061346523],"category_scores_gemma":[0.008736622,0.0008009234,0.0018277246,0.0027960967,0.0014618539,0.006270743,0.0047082873,0.0035718104,0.03323104],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047745518,0.00033263167,0.00110154,0.00080018654,0.00018461648,0.0003532515,0.00050852814,0.0008948631,0.007056717,0.009042729,0.40717167,0.5720757],"study_design_scores_gemma":[0.000060049548,0.00014911343,0.0021606707,0.00057799695,0.00013249162,0.00080394495,0.0003326087,0.005449604,0.0028448163,0.015462432,0.9719678,0.000058407415],"about_ca_topic_score_codex":0.0028930523,"about_ca_topic_score_gemma":0.006930377,"teacher_disagreement_score":0.061346523,"about_ca_system_score_codex":0.0015110048,"about_ca_system_score_gemma":0.0031817213,"threshold_uncertainty_score":0.20522457},"labels":[],"label_agreement":null},{"id":"W2183131996","doi":"10.1186/s40697-015-0084-3","title":"A Little (up)SET THEORY: A Philosophical and Psychological Pondering of a Scientist on the State of Our Art","year":2015,"lang":"en","type":"article","venue":"Canadian Journal of Kidney Health and Disease","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; St. Michael's Hospital","funders":"","keywords":"Zeitgeist; Epistemology; Context (archaeology); Criticism; Paradigm shift; Sociology of scientific knowledge; Cliché; Misnomer; Scientific progress; Cognitive science; Sociology; Psychology; Philosophy; Political science; History; Law","score_opus":0.06757312182026858,"score_gpt":0.3379317211561038,"score_spread":0.2703585993358352,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2183131996","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010237826,0.028795533,0.028324468,0.82159364,0.008799848,0.000035214078,0.00006730503,0.00013690202,0.102009274],"genre_scores_gemma":[0.65760255,0.020902297,0.0342926,0.24314646,0.013178776,0.00028470682,0.00012910418,0.00066871126,0.029794853],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97856265,0.014988972,0.00044429957,0.0016052602,0.00361503,0.00078378257],"domain_scores_gemma":[0.9608267,0.028031707,0.0010708052,0.0033714708,0.00451889,0.0021803277],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.03028632,0.00088263914,0.0016492413,0.003957089,0.01628024,0.020577386,0.0026981167,0.011802049,0.0043820245],"category_scores_gemma":[0.033807393,0.0006910586,0.0011451129,0.0024465304,0.14008702,0.043239415,0.0090190135,0.031914163,0.0015539164],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000024002398,0.000022493829,0.00009971343,0.00005073323,0.000008632016,0.00005392377,0.013898552,0.00008993327,0.00007379522,0.9639724,0.016413521,0.005292378],"study_design_scores_gemma":[0.00001532047,0.000018603847,0.000055659817,0.00013559735,0.00000515437,0.00006741141,0.0054566213,0.00029073076,0.00006660915,0.93123287,0.062631436,0.000023930155],"about_ca_topic_score_codex":0.0043435,"about_ca_topic_score_gemma":0.0032854755,"teacher_disagreement_score":0.98371977,"about_ca_system_score_codex":0.008064848,"about_ca_system_score_gemma":0.0077538677,"threshold_uncertainty_score":0.16017133},"labels":[],"label_agreement":null},{"id":"W2186843366","doi":"","title":"Using Lexical Chaining to Rank Protein-Protein Interactions in Biomedical Text","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Chaining; Computer science; Rank (graph theory); Context (archaeology); Natural language processing; Ranking (information retrieval); Computational linguistics; Artificial intelligence; Information retrieval; Linguistics; Mathematics; Biology; Psychology","score_opus":0.049745427524800136,"score_gpt":0.35089349090727784,"score_spread":0.3011480633824777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2186843366","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.42808345,0.005616807,0.53550816,0.001864141,0.00048564558,0.00092469057,0.010146749,0.0031026958,0.014267685],"genre_scores_gemma":[0.57996935,0.0016675149,0.40416533,0.00019692915,0.00048596738,0.00052579684,0.010478422,0.00019908599,0.0023116139],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99668664,0.0010264586,0.0006987781,0.0005027576,0.00095419236,0.0001312458],"domain_scores_gemma":[0.96969193,0.02295391,0.0026135128,0.00087540335,0.0033099155,0.00055533147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035057433,0.0007438996,0.0007739489,0.021062132,0.0012857425,0.002935164,0.0006758706,0.0010539503,0.005228447],"category_scores_gemma":[0.02243034,0.00025146233,0.00058307644,0.011410569,0.0012759824,0.004137,0.0018025468,0.00069843465,0.0025466017],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009957452,0.00043072368,0.07508202,0.0026754232,0.00046138745,0.001839349,0.0024961098,0.006869185,0.06817202,0.0186229,0.010992702,0.81136245],"study_design_scores_gemma":[0.0005178914,0.0016284991,0.09610634,0.0011931848,0.0012653901,0.004861068,0.008480521,0.41564056,0.13901646,0.27378187,0.056939136,0.00056918914],"about_ca_topic_score_codex":0.001361197,"about_ca_topic_score_gemma":0.0026074986,"teacher_disagreement_score":0.021062132,"about_ca_system_score_codex":0.000525757,"about_ca_system_score_gemma":0.0011394081,"threshold_uncertainty_score":0.018540323},"labels":[],"label_agreement":null},{"id":"W2189229433","doi":"","title":"York University at TREC 2011: Medical Records Track","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Computer science; Query expansion; Information retrieval; Unified Medical Language System; Search engine indexing; Ranking (information retrieval); Matching (statistics); Relevance feedback; Relevance (law); Weighting; Semantic matching; Web search query; Query language; Image retrieval; Search engine; Artificial intelligence; Medicine","score_opus":0.03467695990857757,"score_gpt":0.23252963406210017,"score_spread":0.1978526741535226,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2189229433","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019634295,0.014439569,0.018215312,0.06573671,0.008030111,0.0035658116,0.7451655,0.025536787,0.09967584],"genre_scores_gemma":[0.023019666,0.0044006193,0.018568052,0.005049051,0.0012888106,0.0016620295,0.883591,0.0011976945,0.061222997],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99002236,0.0031785965,0.00079456455,0.0008592062,0.0044223443,0.0007230413],"domain_scores_gemma":[0.9628198,0.009257139,0.0017129285,0.00460683,0.018138189,0.003465017],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017587995,0.0024012534,0.00281536,0.006781226,0.00425907,0.005622849,0.0033295853,0.0037257953,0.05984994],"category_scores_gemma":[0.03250564,0.0008975348,0.0011731611,0.007602286,0.0011268106,0.007954733,0.0021832201,0.0045409356,0.039746657],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010638493,0.00010670876,0.00032632865,0.0002626386,0.000027350688,0.00003035824,0.00003590966,0.00041295023,0.00066277006,0.0007656173,0.9866757,0.0105874315],"study_design_scores_gemma":[0.00074808276,0.00045586756,0.010995728,0.0004151652,0.0001854303,0.0003183724,0.00032998683,0.015911175,0.0063237133,0.0047883624,0.9592635,0.00026452282],"about_ca_topic_score_codex":0.13876764,"about_ca_topic_score_gemma":0.20700984,"teacher_disagreement_score":0.13876764,"about_ca_system_score_codex":0.009194215,"about_ca_system_score_gemma":0.0143153025,"threshold_uncertainty_score":0.27591985},"labels":[],"label_agreement":null},{"id":"W2194609845","doi":"10.6000/1927-5129.2015.11.83","title":"Gene Ontology Tools: A Comparative Study","year":2015,"lang":"en","type":"article","venue":"Journal of Basic & Applied Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Annotation; Ontology; Visualization; Key (lock); Resource (disambiguation); Gene ontology; Information retrieval; Gene Annotation; Data science; World Wide Web; Data mining; Genome; Gene; Artificial intelligence","score_opus":0.1331072111962279,"score_gpt":0.35938478520361117,"score_spread":0.22627757400738327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2194609845","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61636496,0.08400326,0.118487306,0.0029276642,0.0005520306,0.00049015944,0.035409037,0.0033401102,0.13842554],"genre_scores_gemma":[0.808092,0.03574674,0.11415049,0.0004261682,0.00013285904,0.00051211566,0.033474647,0.0009156265,0.006549437],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9977812,0.0006064301,0.00023679054,0.00035910678,0.00082815136,0.000188304],"domain_scores_gemma":[0.99468154,0.0031486212,0.00041688932,0.00037862072,0.001080813,0.00029352953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039156773,0.00040737193,0.0005119411,0.014001705,0.001097366,0.0022021846,0.00072287465,0.0004096524,0.0063889623],"category_scores_gemma":[0.009529576,0.00014445285,0.00089197024,0.017395703,0.000571188,0.0027977591,0.0018022355,0.0004933077,0.0013798145],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019619311,0.00034385434,0.09738269,0.0072074686,0.00108705,0.001578049,0.0072013335,0.00089169975,0.0256007,0.04168512,0.021379042,0.7936811],"study_design_scores_gemma":[0.00009320229,0.0005921313,0.37124047,0.0028911766,0.00195936,0.0065831407,0.01190416,0.0056837914,0.012982056,0.035017155,0.55092037,0.00013302584],"about_ca_topic_score_codex":0.001979598,"about_ca_topic_score_gemma":0.003219669,"teacher_disagreement_score":0.014001705,"about_ca_system_score_codex":0.0010808583,"about_ca_system_score_gemma":0.0013752622,"threshold_uncertainty_score":0.021373212},"labels":[],"label_agreement":null},{"id":"W2196685335","doi":"10.6084/m9.figshare.939458.v1","title":"PhenomeCentral: An Integrated Portal for Sharing and Searching Patient Phenotype Data for Rare Genetic Disorders.","year":2014,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Phenotype; Computer science; Medicine; Computational biology; Genetics; Biology; Gene","score_opus":0.04659821970949185,"score_gpt":0.30045302329085016,"score_spread":0.2538548035813583,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2196685335","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013363624,0.0033347132,0.14196235,0.004554605,0.00065031444,0.0015328478,0.5485543,0.26173714,0.024310151],"genre_scores_gemma":[0.06587553,0.0032644002,0.12859288,0.0027322336,0.0003904005,0.0018962339,0.761013,0.024512185,0.011723178],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99676263,0.0008983636,0.0005036538,0.00067576277,0.0009095558,0.00024987324],"domain_scores_gemma":[0.9834532,0.0053927028,0.0017534035,0.005501751,0.0012108198,0.0026881136],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0076042353,0.0017495696,0.0014831391,0.006364656,0.0008074555,0.0031288809,0.004232134,0.0016516447,0.04628653],"category_scores_gemma":[0.02297498,0.0011108557,0.001061084,0.006177684,0.0006723355,0.0064985813,0.011922114,0.0020839542,0.02953233],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0033850973,0.00021948644,0.012902943,0.0021929345,0.00050469895,0.0016539425,0.0013158742,0.0016872054,0.0066184453,0.011131712,0.81720513,0.1411826],"study_design_scores_gemma":[0.0010032338,0.00026906998,0.021037545,0.0009028885,0.0002585608,0.002606286,0.0007511678,0.011175729,0.009841955,0.031599354,0.9202405,0.0003137675],"about_ca_topic_score_codex":0.0031032306,"about_ca_topic_score_gemma":0.0043391404,"teacher_disagreement_score":0.9957679,"about_ca_system_score_codex":0.0011158112,"about_ca_system_score_gemma":0.0035516198,"threshold_uncertainty_score":0.15484387},"labels":[],"label_agreement":null},{"id":"W2205683799","doi":"10.4018/978-1-60566-208-4.ch006","title":"Social Cognitive Ontology and User Driven Healthcare","year":2009,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"NOSM University","funders":"","keywords":"Ontology; Health care; Premise; Knowledge management; Meaning (existential); Cognition; Key (lock); Feeling; Process (computing); Computer science; Psychology; Epistemology; Social psychology; Psychotherapist; Political science","score_opus":0.024105278095899508,"score_gpt":0.3010555811137719,"score_spread":0.2769503030178724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2205683799","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02439514,0.013599587,0.4962888,0.0698317,0.0012537439,0.0003122122,0.00028663006,0.0006545132,0.39337763],"genre_scores_gemma":[0.8241762,0.007683561,0.13399202,0.0055169985,0.0008794234,0.00084213354,0.00034214288,0.0002440286,0.026323529],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9930074,0.004458166,0.0003236417,0.0005101738,0.001226753,0.00047385518],"domain_scores_gemma":[0.9935314,0.0046439483,0.00030700606,0.0007111171,0.00046320164,0.00034329464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073232222,0.0008106067,0.00064451643,0.0027176957,0.0029256092,0.009489728,0.0019946517,0.0040605557,0.005135034],"category_scores_gemma":[0.0068443837,0.0005733512,0.0012499071,0.0024174913,0.02858799,0.010539076,0.0065020267,0.004167202,0.000568279],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000035163164,0.0000094961015,0.00008290434,0.000036249843,0.0000050062554,0.0000577535,0.0036206918,0.00047441333,0.000050393723,0.9917292,0.0008057177,0.0031245938],"study_design_scores_gemma":[0.000007486182,0.000011665612,0.00014400348,0.00008406068,0.0000056810504,0.000118536074,0.0025643539,0.0033879557,0.00011257856,0.9409702,0.052578367,0.000015129357],"about_ca_topic_score_codex":0.0067845094,"about_ca_topic_score_gemma":0.003264579,"teacher_disagreement_score":0.009489728,"about_ca_system_score_codex":0.008053455,"about_ca_system_score_gemma":0.0042354995,"threshold_uncertainty_score":0.05843216},"labels":[],"label_agreement":null},{"id":"W2210299414","doi":"10.1007/978-3-642-13131-8_1","title":"Overview of the Ninth Annual Meeting of the BioLINK SIG at ISMB: Linking Literature, Information and Knowledge for Biology","year":2010,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Automatic summarization; Computer science; Information extraction; Identification (biology); Knowledge extraction; Domain (mathematical analysis); Data science; Information retrieval; Key (lock); Artificial intelligence; Biology; Ecology","score_opus":0.01710943575321675,"score_gpt":0.27772978818555255,"score_spread":0.2606203524323358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2210299414","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050375704,0.32721692,0.089991875,0.121034585,0.19989523,0.0024114992,0.020769438,0.006109881,0.2275329],"genre_scores_gemma":[0.0081452215,0.1296164,0.06631972,0.016509056,0.02638722,0.0010982666,0.0316191,0.0023925586,0.7179125],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985044,0.00025267855,0.00011780277,0.00026647176,0.00068682135,0.00017176988],"domain_scores_gemma":[0.99646854,0.00028917813,0.00013980555,0.00014567966,0.0016516038,0.0013051737],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.004926265,0.0018140095,0.0015702223,0.0064128414,0.0015540237,0.007084649,0.0018992128,0.0022073246,0.08048338],"category_scores_gemma":[0.0034425799,0.0007385398,0.0011681564,0.006010392,0.0005869522,0.005079724,0.004110445,0.0023724977,0.057148464],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010621067,0.00008933189,0.0002201366,0.0005556363,0.000019400113,0.000077036515,0.00007947047,0.00019262868,0.0016477795,0.0013950795,0.79841506,0.1972022],"study_design_scores_gemma":[0.000008379335,0.000029790019,0.00043322673,0.00020326862,0.000012581318,0.00007823672,0.000044980556,0.00014825781,0.00038182535,0.0008099598,0.99783474,0.00001483249],"about_ca_topic_score_codex":0.0047634053,"about_ca_topic_score_gemma":0.0155763915,"teacher_disagreement_score":0.99291533,"about_ca_system_score_codex":0.00192491,"about_ca_system_score_gemma":0.0051328107,"threshold_uncertainty_score":0.26924372},"labels":[],"label_agreement":null},{"id":"W2211589933","doi":"10.1007/978-3-642-36089-3_1","title":"Addressing Cognitive and Social Challenges in Designing and Using Ontologies in the Biomedical Domain","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Conceptualization; Ontology; Interoperability; Categorization; Domain (mathematical analysis); Data science; Vocabulary; Reuse; Semantic interoperability; IDEF5; Information retrieval; Ontology-based data integration; Artificial intelligence; World Wide Web; Epistemology","score_opus":0.10774980924205195,"score_gpt":0.3228602348282431,"score_spread":0.21511042558619115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2211589933","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0134940995,0.0068113455,0.91433185,0.032539938,0.00056295993,0.00017230197,0.00022587818,0.00046734067,0.031394336],"genre_scores_gemma":[0.102855094,0.007619545,0.8780246,0.0027732486,0.0006247731,0.00033441855,0.0006198792,0.00032364836,0.00682494],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98952824,0.005744127,0.0011071356,0.00086904457,0.002335677,0.00041582648],"domain_scores_gemma":[0.9697772,0.024496112,0.0011691883,0.0023549811,0.0015777068,0.00062474445],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02130128,0.0008692138,0.0011866258,0.0036175528,0.0038085268,0.01522386,0.0030003053,0.0034684748,0.0030407798],"category_scores_gemma":[0.029494269,0.0010980013,0.0018478726,0.0046886583,0.00776681,0.03555679,0.008875444,0.0056336084,0.0012250452],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000247051,0.000106641324,0.0009977889,0.0008762874,0.00010479194,0.00038820494,0.00833784,0.0043236497,0.0026461023,0.7741553,0.01201653,0.19602203],"study_design_scores_gemma":[0.0000068773265,0.00001576724,0.00039546686,0.00038469964,0.00004905313,0.00046116655,0.0043778596,0.016335735,0.0016476136,0.86695033,0.10933141,0.000044062792],"about_ca_topic_score_codex":0.0035003552,"about_ca_topic_score_gemma":0.0057935137,"teacher_disagreement_score":0.02130128,"about_ca_system_score_codex":0.0025316097,"about_ca_system_score_gemma":0.0041996627,"threshold_uncertainty_score":0.112653315},"labels":[],"label_agreement":null},{"id":"W2213189003","doi":"","title":"Chemistry-specific Features and Heuristics for Developing a CRF-based Chemical Named Entity Recogniser","year":2013,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"Wellcome Trust","keywords":"Conditional random field; Named-entity recognition; Heuristics; Computer science; Artificial intelligence; Task (project management); Natural language processing; Training set; Set (abstract data type); Sequence labeling; Security token; Search engine indexing; Pattern recognition (psychology); Machine learning","score_opus":0.06544131363573949,"score_gpt":0.2888142198452974,"score_spread":0.22337290620955794,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2213189003","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025972586,0.000552375,0.9379559,0.00024100611,0.00013170125,0.0005610115,0.0031534622,0.028422624,0.0030092627],"genre_scores_gemma":[0.12305185,0.00026582234,0.8639789,0.00021619386,0.000058808037,0.0005131043,0.008554676,0.0008327663,0.002527877],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99887806,0.00027125477,0.00013905126,0.00039604146,0.00020486614,0.00011071005],"domain_scores_gemma":[0.99691904,0.0020168421,0.000162213,0.00037832223,0.00043487226,0.000088740286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025504236,0.001275195,0.0009797319,0.0014485926,0.0005129221,0.0010390293,0.0019068288,0.001408004,0.0069837375],"category_scores_gemma":[0.0051521095,0.00066829904,0.0009027188,0.0011836039,0.0005206207,0.0021090894,0.0008551898,0.0015051349,0.006442994],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006450445,0.00041994243,0.0038578669,0.0010024837,0.00019527425,0.00062382844,0.000306034,0.10906723,0.07879466,0.0059420033,0.020691814,0.7784538],"study_design_scores_gemma":[0.00018475692,0.00050523964,0.004058304,0.00009660976,0.00021640229,0.0009143968,0.00022642461,0.84009707,0.11078564,0.008888615,0.033850573,0.0001758798],"about_ca_topic_score_codex":0.005072605,"about_ca_topic_score_gemma":0.008851959,"teacher_disagreement_score":0.0069837375,"about_ca_system_score_codex":0.00071098364,"about_ca_system_score_gemma":0.0014931365,"threshold_uncertainty_score":0.023362875},"labels":[],"label_agreement":null},{"id":"W222181838","doi":"10.1186/s12859-015-0453-z","title":"MeSH ORA framework: R/Bioconductor packages to support MeSH over-representation analysis","year":2015,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute of Genetics; RIKEN; Research Organization of Information and Systems","keywords":"Computer science","score_opus":0.0652034063736068,"score_gpt":0.3488638436000229,"score_spread":0.2836604372264161,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W222181838","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022277262,0.0017683101,0.40258223,0.001387138,0.0014313234,0.0007422217,0.11695118,0.4640723,0.008837546],"genre_scores_gemma":[0.027951762,0.0018657345,0.6673913,0.0019247261,0.00050148007,0.00887402,0.122746184,0.16022046,0.008524341],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99539906,0.0013010069,0.00063421857,0.0013324387,0.0010038887,0.00032942012],"domain_scores_gemma":[0.9851247,0.008011955,0.001561387,0.002477126,0.002210724,0.0006139895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009416809,0.0046961787,0.003813253,0.009706597,0.0016607391,0.004338349,0.008475476,0.002344449,0.08993926],"category_scores_gemma":[0.037796836,0.002625742,0.0047617736,0.008365436,0.0012843736,0.0038591085,0.0061073694,0.003820694,0.060265336],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013170523,0.00012117083,0.0039251633,0.009739837,0.0026117568,0.00059249636,0.000886708,0.0051661357,0.00678894,0.033208273,0.84820384,0.08743866],"study_design_scores_gemma":[0.0007735131,0.00020689101,0.0077867997,0.0017929489,0.0015970925,0.0009271457,0.00028128512,0.05520011,0.010978715,0.09751735,0.8225274,0.00041073113],"about_ca_topic_score_codex":0.0053146095,"about_ca_topic_score_gemma":0.0054663494,"teacher_disagreement_score":0.08993926,"about_ca_system_score_codex":0.001440737,"about_ca_system_score_gemma":0.0051177144,"threshold_uncertainty_score":0.3008768},"labels":[],"label_agreement":null},{"id":"W2224038009","doi":"10.1007/978-3-319-25010-6_16","title":"Automatic Curation of Clinical Trials Data in LinkedCT","year":2015,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; XML; Raw data; Data quality; Quality (philosophy); Clinical trial; Source document; Data curation; Linked data; Data science; World Wide Web; Information retrieval; Database; Bioinformatics; Semantic Web","score_opus":0.2551144569738897,"score_gpt":0.4541076386840874,"score_spread":0.1989931817101977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2224038009","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016466321,0.014083708,0.5962127,0.004051704,0.0013865179,0.0026961663,0.2996827,0.0488088,0.016611421],"genre_scores_gemma":[0.046239346,0.006218303,0.6606217,0.0015379108,0.00042703812,0.0015509129,0.27347222,0.003962295,0.0059702396],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99285805,0.0024835274,0.0014876409,0.0011003983,0.0018458985,0.00022449846],"domain_scores_gemma":[0.96254027,0.023355803,0.0030055866,0.007090939,0.0033112224,0.00069613016],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010454593,0.0014245354,0.0023265965,0.01941641,0.0014865204,0.0068905875,0.0020641545,0.0014852753,0.011625496],"category_scores_gemma":[0.04501898,0.0011306631,0.0039003822,0.013349601,0.00088165654,0.003729923,0.0061900117,0.0025034607,0.00685567],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051870657,0.00032711076,0.011143444,0.014356266,0.0014877466,0.0015051119,0.002282309,0.013101043,0.015652806,0.054359,0.25129458,0.6339718],"study_design_scores_gemma":[0.00025229398,0.00018558037,0.012947273,0.0054573123,0.001779465,0.001979024,0.00054825266,0.053722132,0.032083232,0.1262112,0.7645824,0.00025184255],"about_ca_topic_score_codex":0.007611115,"about_ca_topic_score_gemma":0.021275945,"teacher_disagreement_score":0.9895454,"about_ca_system_score_codex":0.0022706804,"about_ca_system_score_gemma":0.009917775,"threshold_uncertainty_score":0.055289865},"labels":[],"label_agreement":null},{"id":"W2237301598","doi":"","title":"A showcase of Health Science related courses, resources and training materials developed at the Centre for e-Learning at the University of Ottawa","year":2007,"lang":"en","type":"article","venue":"EdMedia: World Conference on Educational Media and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Ottawa","funders":"","keywords":"Training (meteorology); Medical education; Library science; Medicine; Computer science; Geography","score_opus":0.03088464680146061,"score_gpt":0.28301744370884296,"score_spread":0.25213279690738233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2237301598","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3572264,0.0019519293,0.03511391,0.034888588,0.006341176,0.0027323759,0.02778252,0.013011087,0.52095205],"genre_scores_gemma":[0.36344516,0.0016496277,0.04293129,0.0028575163,0.0010984772,0.00074709905,0.015147965,0.0036812718,0.5684417],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99882656,0.000313002,0.000022530709,0.000120923476,0.00030919854,0.0004078215],"domain_scores_gemma":[0.99340653,0.0020994062,0.00009538147,0.00031672165,0.0004646332,0.0036174613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018699629,0.0014502718,0.0005847388,0.001717989,0.0055086655,0.0032628805,0.0020231467,0.0029785268,0.09965823],"category_scores_gemma":[0.003278106,0.00066798413,0.0008570069,0.0020406554,0.0013763562,0.0020555416,0.004604663,0.0029155007,0.01640438],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002451258,0.003786482,0.0051113656,0.0016024841,0.000071201845,0.007482791,0.029854363,0.0043480345,0.034247294,0.007812987,0.76710975,0.13612196],"study_design_scores_gemma":[0.00016512543,0.00065625226,0.013544857,0.00018233116,0.00003490753,0.0011356575,0.012244965,0.0028179188,0.0073300633,0.001292344,0.96049356,0.00010198242],"about_ca_topic_score_codex":0.050176185,"about_ca_topic_score_gemma":0.21043321,"teacher_disagreement_score":0.09965823,"about_ca_system_score_codex":0.0029821952,"about_ca_system_score_gemma":0.0030322876,"threshold_uncertainty_score":0.33338994},"labels":[],"label_agreement":null},{"id":"W2248009881","doi":"10.1109/smc.2015.332","title":"Flexible Concept Matching for Medical Information Retrieval","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Matching (statistics); Flexibility (engineering); Information retrieval; Synonym (taxonomy); Identification (biology); Process (computing); Phrase; Artificial intelligence; Domain (mathematical analysis); Data mining; Natural language processing; Mathematics","score_opus":0.02578924653828432,"score_gpt":0.31160635380380775,"score_spread":0.28581710726552345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2248009881","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018355895,0.0026452611,0.972061,0.00045847,0.00010856663,0.00048136312,0.0005369266,0.0025058053,0.0028466901],"genre_scores_gemma":[0.14293315,0.0011558344,0.8518864,0.0003013142,0.00011603398,0.00029682866,0.001460649,0.00014133251,0.0017085047],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99371344,0.002318141,0.0008084431,0.00088524783,0.0020501672,0.00022451239],"domain_scores_gemma":[0.99572396,0.0020454505,0.0004041421,0.0011546122,0.00053128036,0.00014049376],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005226803,0.00059138105,0.0012105372,0.0065482487,0.0008695969,0.0018214675,0.0018409088,0.0012405644,0.0029984901],"category_scores_gemma":[0.013271647,0.00040157555,0.0012690605,0.007644622,0.0010466344,0.005215683,0.0027744595,0.0009403066,0.0016942926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028181763,0.000219271,0.0015768689,0.00061219354,0.00014552372,0.0003205165,0.0004720838,0.016148113,0.021440245,0.051510505,0.0085424455,0.89873046],"study_design_scores_gemma":[0.00019008019,0.00046578754,0.0041158525,0.0002570468,0.00023758347,0.0027688213,0.0007123285,0.46352503,0.045755144,0.37804523,0.10369153,0.00023557816],"about_ca_topic_score_codex":0.0024473497,"about_ca_topic_score_gemma":0.0017884985,"teacher_disagreement_score":0.0065482487,"about_ca_system_score_codex":0.0010717564,"about_ca_system_score_gemma":0.0015785986,"threshold_uncertainty_score":0.02764231},"labels":[],"label_agreement":null},{"id":"W2248409046","doi":"10.1522/18358157","title":"hyperglycerolemie familiale au Saguenay-Lac-Saint-Jean :","year":2004,"lang":"fr","type":"book","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"SAINT; Geography; History; Art history","score_opus":0.018914841191461876,"score_gpt":0.25229425018857715,"score_spread":0.2333794089971153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2248409046","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.998796,0.00016028818,0.00028966894,0.000044570017,0.000007447765,0.00000969505,0.0000887245,0.0000091674365,0.00059435447],"genre_scores_gemma":[0.9969009,0.00025970675,0.0006471275,0.000039771152,0.0000067075293,0.0000102595295,0.00011254364,0.0000081221815,0.002014887],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.999778,0.000045470777,0.000014143049,0.0000662452,0.000058236375,0.000037814556],"domain_scores_gemma":[0.9997094,0.000078193174,0.00006822393,0.000017280618,0.00006634466,0.00006056943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00026162472,0.0004586444,0.0002825958,0.00073447486,0.0010999722,0.00055164564,0.00020626784,0.00051945465,0.0024759984],"category_scores_gemma":[0.00072329154,0.00017028935,0.00029092203,0.0006269902,0.00052346324,0.00012706981,0.0003648566,0.000393385,0.00024785195],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008320578,0.00015166112,0.82782626,0.00014573276,0.00018122609,0.04010838,0.010319086,0.00045247818,0.09520561,0.0007249649,0.0004999625,0.02355251],"study_design_scores_gemma":[0.000018427909,0.00042474057,0.9343386,0.00004771274,0.000090203546,0.04479539,0.0029910451,0.0007261948,0.0088471975,0.00013662779,0.0075497027,0.000034064917],"about_ca_topic_score_codex":0.0939329,"about_ca_topic_score_gemma":0.067604,"teacher_disagreement_score":0.90606713,"about_ca_system_score_codex":0.000687209,"about_ca_system_score_gemma":0.0007159611,"threshold_uncertainty_score":0.18677229},"labels":[],"label_agreement":null},{"id":"W2249465597","doi":"10.5281/zenodo.27369","title":"plann: APPS publication version","year":2015,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Internet privacy; World Wide Web; Business; Information retrieval","score_opus":0.05111145706368675,"score_gpt":0.26609432534838273,"score_spread":0.21498286828469598,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2249465597","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013023044,0.00041678984,0.043372467,0.0007272376,0.0009034815,0.0003245913,0.49406108,0.40375575,0.05513626],"genre_scores_gemma":[0.011350956,0.00061638706,0.04661789,0.0008507731,0.00031402672,0.0010627283,0.6939529,0.16599195,0.07924242],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999278,0.00007111104,0.00007831045,0.00018609702,0.00029228767,0.0000942079],"domain_scores_gemma":[0.9980312,0.0006780688,0.00007803823,0.0004578119,0.00061316765,0.0001416636],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011039743,0.002482364,0.0012194506,0.002721327,0.00068853586,0.002959553,0.002296926,0.0016147239,0.47457957],"category_scores_gemma":[0.0057799094,0.0015765098,0.0013031816,0.0021852446,0.0004538011,0.0033710024,0.0027469571,0.0016755827,0.39928788],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021531933,0.00002485092,0.00025359594,0.0004368241,0.000017750654,0.00007565133,0.000059698916,0.00032372773,0.00074046635,0.0011561052,0.96542484,0.031271126],"study_design_scores_gemma":[0.00020391389,0.000024070077,0.00107243,0.00022138677,0.000023092582,0.00015693103,0.00004714068,0.003573042,0.0034372346,0.0076796655,0.98347694,0.00008421754],"about_ca_topic_score_codex":0.005416494,"about_ca_topic_score_gemma":0.006341905,"teacher_disagreement_score":0.47457957,"about_ca_system_score_codex":0.00083641143,"about_ca_system_score_gemma":0.0013366484,"threshold_uncertainty_score":0.7494484},"labels":[],"label_agreement":null},{"id":"W2250597415","doi":"10.63317/5k96z93j8umc","title":"The Meta-knowledge of Causality in Biomedical Scientific Discourse","year":2014,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"Engineering and Physical Sciences Research Council; Medical Research Council","keywords":"Causality (physics); Computer science; Context (archaeology); Natural language processing; Workload; Artificial intelligence; Data science; Information retrieval; Biology","score_opus":0.03575341318553831,"score_gpt":0.33919154765902604,"score_spread":0.3034381344734877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250597415","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08461027,0.02338717,0.82047623,0.03667357,0.0010061723,0.00025820764,0.0042401245,0.0007546454,0.028593656],"genre_scores_gemma":[0.76895165,0.009051866,0.21204081,0.0023057505,0.0013186549,0.0003837797,0.003222048,0.00020022363,0.0025251752],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9736475,0.012975324,0.0047203447,0.0037977393,0.00439277,0.00046629936],"domain_scores_gemma":[0.8623863,0.11391181,0.00893305,0.008902128,0.0046847635,0.001181943],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025805006,0.001052189,0.0016502292,0.020647109,0.0030501469,0.013861426,0.0029309525,0.0034129284,0.004731752],"category_scores_gemma":[0.09003053,0.0019731158,0.0026130353,0.01409691,0.009205137,0.047026455,0.0065075643,0.0041983686,0.00064786547],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015906703,0.000067174245,0.008833406,0.0014441998,0.0004322919,0.0009165713,0.006378146,0.003154573,0.0014471903,0.9161795,0.0019607318,0.059027202],"study_design_scores_gemma":[0.000022952903,0.000016694183,0.001677105,0.0006299676,0.00030973984,0.00051434874,0.0011346649,0.0070188516,0.0009795644,0.97716296,0.010495286,0.00003778178],"about_ca_topic_score_codex":0.0032180117,"about_ca_topic_score_gemma":0.0024573863,"teacher_disagreement_score":0.025805006,"about_ca_system_score_codex":0.0031583884,"about_ca_system_score_gemma":0.00460364,"threshold_uncertainty_score":0.13647157},"labels":[],"label_agreement":null},{"id":"W2250644920","doi":"","title":"Building a Patient-based Ontology for User-written Web Messages","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Ontology; Computer science; Vocabulary; World Wide Web; OWL-S; Information retrieval; Ontology-based data integration; Semantic Web; Upper ontology; Social Semantic Web; Linguistics","score_opus":0.02988730453320246,"score_gpt":0.2754442116652726,"score_spread":0.24555690713207015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250644920","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011482107,0.00011439448,0.96371,0.0014993662,0.00018104298,0.0013250866,0.007516995,0.0077327415,0.006438208],"genre_scores_gemma":[0.08040267,0.0003056974,0.8985149,0.0005801295,0.00007683279,0.00085668964,0.013090805,0.0011326689,0.0050395476],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9952136,0.0007660932,0.0010365428,0.00093005877,0.0017076374,0.00034592836],"domain_scores_gemma":[0.99434,0.0016728875,0.0006003857,0.0012683703,0.0017567853,0.00036163742],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005681525,0.0008496346,0.0010113028,0.004826198,0.0020234114,0.0046241996,0.0022704103,0.0018456515,0.004881742],"category_scores_gemma":[0.010557584,0.0010263802,0.002666139,0.003745152,0.0012523052,0.008260501,0.0036651255,0.0026350177,0.0025079383],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043315036,0.0012224304,0.028371675,0.0020275556,0.00047827687,0.0038627868,0.0124238795,0.04157737,0.03293199,0.35775894,0.06485291,0.45405897],"study_design_scores_gemma":[0.00013786493,0.00018104246,0.009289497,0.000955289,0.00055097864,0.0024441455,0.003100309,0.27908465,0.040635053,0.13439542,0.528841,0.00038469135],"about_ca_topic_score_codex":0.024868969,"about_ca_topic_score_gemma":0.028559506,"teacher_disagreement_score":0.024868969,"about_ca_system_score_codex":0.0033398517,"about_ca_system_score_gemma":0.006821664,"threshold_uncertainty_score":0.04944843},"labels":[],"label_agreement":null},{"id":"W2250670243","doi":"","title":"A Machine Learning Approach for Phenotype Name Recognition","year":2012,"lang":"en","type":"article","venue":"Scholarship@Western (Western University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Artificial intelligence; Natural language processing; Process (computing); Machine learning; Named-entity recognition; Phenotype; Engineering; Gene; Biology","score_opus":0.10951118899604961,"score_gpt":0.3143242810529213,"score_spread":0.2048130920568717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2250670243","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013258647,0.0007933678,0.9708046,0.0010680249,0.00024205097,0.00032339568,0.0018660101,0.007903475,0.0037405025],"genre_scores_gemma":[0.085325256,0.000438208,0.90666825,0.00040481996,0.00016063575,0.0003334423,0.0039476766,0.00013085276,0.002590865],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978131,0.00046763808,0.00030700176,0.0007669135,0.0005756713,0.00006958274],"domain_scores_gemma":[0.99563843,0.0023996758,0.0003876801,0.0004896144,0.0009832174,0.00010143599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021833137,0.00082296884,0.00085705257,0.004529595,0.00095699483,0.0018652555,0.001746397,0.0013438783,0.0022852567],"category_scores_gemma":[0.008747719,0.0003148144,0.0010240366,0.0031209013,0.00060841563,0.0024636842,0.0009619215,0.00138901,0.002424996],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013590501,0.00036935904,0.0060906042,0.0003596504,0.00013106003,0.000640503,0.00031894623,0.010252476,0.015830666,0.010086074,0.014004763,0.9417799],"study_design_scores_gemma":[0.000090329595,0.00027494182,0.009136161,0.00021845082,0.00024022395,0.0028002681,0.00043554354,0.79398614,0.04190312,0.08453715,0.066201225,0.00017651723],"about_ca_topic_score_codex":0.0020017608,"about_ca_topic_score_gemma":0.0026766378,"teacher_disagreement_score":0.004529595,"about_ca_system_score_codex":0.00069438753,"about_ca_system_score_gemma":0.001382723,"threshold_uncertainty_score":0.011546612},"labels":[],"label_agreement":null},{"id":"W2251261672","doi":"","title":"Extracting Information for Generating A Diabetes Report Card from Free Text in Physicians Notes","year":2010,"lang":"en","type":"article","venue":"North American Chapter of the Association for Computational Linguistics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada; McMaster University; University of Ottawa","funders":"","keywords":"Text messaging; Computer science; Guideline; Diabetes mellitus; Population; Health records; Process (computing); Information retrieval; Medicine; Data mining; World Wide Web","score_opus":0.009028242911674093,"score_gpt":0.255037806892002,"score_spread":0.24600956398032792,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2251261672","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19668855,0.0017342462,0.57461166,0.0036984796,0.00071177457,0.0046044127,0.16605113,0.04265761,0.009242099],"genre_scores_gemma":[0.11931767,0.00075103843,0.72792774,0.00036723402,0.00018160768,0.001056766,0.14746438,0.00043110663,0.002502471],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99789125,0.00043417737,0.0004339948,0.00048321608,0.0006474072,0.00011001619],"domain_scores_gemma":[0.9882108,0.007699793,0.0010785687,0.0010773316,0.00171291,0.00022067346],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023044925,0.0012276412,0.0008690779,0.0071838065,0.00075960904,0.0020825537,0.0010652452,0.0014459436,0.004174344],"category_scores_gemma":[0.014562051,0.00047598657,0.0011125173,0.004019812,0.00036865677,0.002370662,0.0012264387,0.0010379906,0.0059876367],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082227075,0.0008147187,0.025630923,0.0016275125,0.00015307748,0.0016759452,0.0011940374,0.008630565,0.027515968,0.003515142,0.06753212,0.8608877],"study_design_scores_gemma":[0.0006826077,0.0008276475,0.07426293,0.0012860647,0.00072673534,0.003914998,0.0058106487,0.42584375,0.20952117,0.027969958,0.24874021,0.00041329858],"about_ca_topic_score_codex":0.004792041,"about_ca_topic_score_gemma":0.0053333184,"teacher_disagreement_score":0.0071838065,"about_ca_system_score_codex":0.000937253,"about_ca_system_score_gemma":0.0026357071,"threshold_uncertainty_score":0.013964593},"labels":[],"label_agreement":null},{"id":"W2252897399","doi":"10.1002/pra2.2015.14505201006","title":"How can information science contribute to alzheimer's disease research?","year":2015,"lang":"en","type":"article","venue":"Proceedings of the Association for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"University of Pennsylvania","keywords":"Session (web analytics); Service (business); Health care; Information exchange; Psychology; Medicine; Political science; Computer science; Business; World Wide Web","score_opus":0.03116937335244972,"score_gpt":0.3118085604749567,"score_spread":0.28063918712250696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2252897399","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043032807,0.057060942,0.009404304,0.8688612,0.0053319116,0.00016943358,0.00018384511,0.000113646594,0.05457146],"genre_scores_gemma":[0.46436465,0.2622024,0.043435607,0.18232295,0.03099962,0.001189975,0.00044043735,0.00020825332,0.0148359975],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.93343455,0.056874406,0.0014029805,0.0012691899,0.0054231896,0.0015958052],"domain_scores_gemma":[0.65167236,0.3165734,0.0037154793,0.007077736,0.014694046,0.0062670293],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11197389,0.0011648968,0.001726911,0.01478208,0.0066772345,0.032834176,0.0023600266,0.007826993,0.010984612],"category_scores_gemma":[0.11582565,0.00065216917,0.0015278916,0.011532002,0.021824364,0.039661117,0.011280855,0.007890955,0.0021164636],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000114082315,0.00023063837,0.003781355,0.0048515415,0.00017678685,0.00031214286,0.015824758,0.00071226823,0.000328903,0.59133005,0.17259483,0.20974262],"study_design_scores_gemma":[0.00006182933,0.00013261898,0.001184704,0.008035347,0.00012530551,0.00022350885,0.021634119,0.0011463735,0.00054782274,0.5027265,0.46410337,0.000078575125],"about_ca_topic_score_codex":0.003649023,"about_ca_topic_score_gemma":0.0029292433,"teacher_disagreement_score":0.11197389,"about_ca_system_score_codex":0.008297968,"about_ca_system_score_gemma":0.017683145,"threshold_uncertainty_score":0.59218156},"labels":[],"label_agreement":null},{"id":"W2260422212","doi":"10.1371/journal.pone.0131274","title":"PhenStat: A Tool Kit for Standardized Analysis of High Throughput Phenotypic Data","year":2015,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":67,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Human Genome Research Institute; National Institutes of Health; Institute of Genetics; Wellcome Trust; Wellcome","keywords":"Throughput; Phenotype; Computational biology; Computer science; Biology; Genetics; Operating system","score_opus":0.14304798998158783,"score_gpt":0.3158650875558744,"score_spread":0.17281709757428657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2260422212","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00580513,0.0010087952,0.42970923,0.00043585943,0.0005291775,0.0013319724,0.2174178,0.33916724,0.0045948187],"genre_scores_gemma":[0.019699506,0.001423408,0.63928235,0.0008583237,0.00024444688,0.012408669,0.2219561,0.09709901,0.007028183],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99375844,0.0016553171,0.0012164214,0.0010534059,0.0019515797,0.00036482816],"domain_scores_gemma":[0.9843078,0.008843089,0.001881537,0.002649371,0.0018360469,0.0004821797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0104980385,0.0028287515,0.0028489719,0.0054028346,0.0009126016,0.0028514927,0.0035509432,0.0009803888,0.052672632],"category_scores_gemma":[0.018744482,0.0023767801,0.0028118307,0.004053468,0.00092533376,0.0027893067,0.003774235,0.0035370658,0.030303016],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018814299,0.00029130708,0.009893943,0.0071994276,0.0015951445,0.00093335554,0.0006589366,0.007914969,0.037863076,0.010675319,0.7445822,0.17651086],"study_design_scores_gemma":[0.0009273867,0.00055473624,0.027859489,0.0011137754,0.0006360604,0.0014952047,0.00025631458,0.037910182,0.053678516,0.037313677,0.83762246,0.0006322015],"about_ca_topic_score_codex":0.0013297442,"about_ca_topic_score_gemma":0.0028117679,"teacher_disagreement_score":0.052672632,"about_ca_system_score_codex":0.0007797434,"about_ca_system_score_gemma":0.003828546,"threshold_uncertainty_score":0.17620748},"labels":[],"label_agreement":null},{"id":"W2260849334","doi":"10.1016/j.jbmt.2015.12.011","title":"Simulation of abstract models of structural homeostasis","year":2016,"lang":"en","type":"editorial","venue":"Journal of Bodywork and Movement Therapies","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Homeostasis; Neuroscience; Medicine; Psychology; Internal medicine","score_opus":0.013177172095055563,"score_gpt":0.2936053813351234,"score_spread":0.28042820924006784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2260849334","genre_codex":"methods","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036139578,0.0022727852,0.81645596,0.060871046,0.046037246,0.00012331665,0.0034388814,0.0041967006,0.030464545],"genre_scores_gemma":[0.69995856,0.005390687,0.2227512,0.0063246144,0.0134592885,0.000704668,0.004632646,0.0024174512,0.04436083],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99931467,0.00028067495,0.00005519749,0.00007646418,0.00023936098,0.000033674172],"domain_scores_gemma":[0.9915731,0.007151529,0.00015449048,0.00031520537,0.00058271026,0.0002229329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015076459,0.0008337739,0.0008786174,0.0006359884,0.00063521945,0.0022880905,0.0019939097,0.0024858285,0.0109252175],"category_scores_gemma":[0.0142442435,0.00053802255,0.0017794744,0.00045595743,0.0011008559,0.0027874399,0.00161395,0.0029350077,0.001098561],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030297003,0.000059126738,0.0005579276,0.0004939958,0.00024103688,0.00050138717,0.00027717298,0.6232618,0.00198335,0.22230808,0.11643369,0.033579443],"study_design_scores_gemma":[0.00016529673,0.000026236683,0.00013118266,0.00006284856,0.000079480385,0.000058845195,0.00004482615,0.77602595,0.0012222223,0.16333294,0.05882012,0.000030075913],"about_ca_topic_score_codex":0.0029118266,"about_ca_topic_score_gemma":0.0027238245,"teacher_disagreement_score":0.0109252175,"about_ca_system_score_codex":0.001273024,"about_ca_system_score_gemma":0.0011009874,"threshold_uncertainty_score":0.036548436},"labels":[],"label_agreement":null},{"id":"W2281601788","doi":"","title":"The Canadian Thyroid Cancer Consortium Registry development experience","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal; Nova Scotia Cancer Centre; London Health Sciences Centre","funders":"","keywords":"Computer science; Database; Identifier; Documentation; World Wide Web; Audit; Data dictionary; Metadata; Information retrieval","score_opus":0.019266473497500457,"score_gpt":0.3003387023930721,"score_spread":0.28107222889557165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2281601788","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36392325,0.0223132,0.019941004,0.1839245,0.003176955,0.010224107,0.05377971,0.004560986,0.33815634],"genre_scores_gemma":[0.753573,0.020326966,0.083117284,0.015938338,0.000736674,0.004905379,0.061653964,0.0018741944,0.057874203],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.9653999,0.00956364,0.0021951944,0.0025533407,0.015496213,0.0047916574],"domain_scores_gemma":[0.91136855,0.011218066,0.0035984113,0.008221417,0.044804614,0.02078894],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07138609,0.0006883368,0.0005171825,0.004820707,0.0055054766,0.0065069506,0.006416814,0.001238113,0.007650163],"category_scores_gemma":[0.069231406,0.0009872909,0.0007651647,0.013805456,0.0021381818,0.0029071372,0.007488384,0.0019549308,0.0017737773],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005296314,0.00051012833,0.18991843,0.0008002848,0.000109544075,0.0022397735,0.016698422,0.0021540534,0.0010809066,0.020993209,0.3696757,0.3952899],"study_design_scores_gemma":[0.00017299736,0.00023251517,0.11809793,0.00088479545,0.00007174396,0.001696669,0.012535594,0.0015387153,0.0009279884,0.0011480418,0.86253846,0.00015459217],"about_ca_topic_score_codex":0.9046868,"about_ca_topic_score_gemma":0.9348532,"teacher_disagreement_score":0.9286139,"about_ca_system_score_codex":0.07295514,"about_ca_system_score_gemma":0.2957536,"threshold_uncertainty_score":0.52932906},"labels":[],"label_agreement":null},{"id":"W2285152544","doi":"","title":"La conceptualisation métaphorique en biomédecine : indices de conceptualisation et réseaux lexicaux","year":2006,"lang":"fr","type":"article","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Cambridge; Université de Montréal; University of Chicago; Sanofi","keywords":"Humanities; Philosophy; Sociology","score_opus":0.0077043711619893505,"score_gpt":0.21306306507497352,"score_spread":0.20535869391298417,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2285152544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16434644,0.02168182,0.70294505,0.010938477,0.00075362175,0.00060832466,0.0026504574,0.001062707,0.09501314],"genre_scores_gemma":[0.6954431,0.0045923125,0.29238135,0.00047545257,0.00015200391,0.0005617067,0.0015552142,0.00027556735,0.004563213],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9910999,0.0052659907,0.00086644694,0.0007509007,0.0017024175,0.00031435204],"domain_scores_gemma":[0.973836,0.01933188,0.0019311054,0.001393976,0.0027872222,0.00071971375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010241644,0.0010159154,0.0007503093,0.016874,0.0027585314,0.01245737,0.0010765155,0.0021691055,0.0057342905],"category_scores_gemma":[0.049674466,0.000702717,0.0010550024,0.021098398,0.008508901,0.029539425,0.00440788,0.0032454655,0.0011599173],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025296735,0.000073160314,0.010070929,0.0012955429,0.000110560584,0.0003167348,0.043464664,0.0015126396,0.003791969,0.81533164,0.005719148,0.11805994],"study_design_scores_gemma":[0.000085611144,0.00018094589,0.013958576,0.00131324,0.00016881089,0.0022269345,0.05646134,0.017240902,0.0033176453,0.7462831,0.15857652,0.00018639797],"about_ca_topic_score_codex":0.00562912,"about_ca_topic_score_gemma":0.003923448,"teacher_disagreement_score":0.016874,"about_ca_system_score_codex":0.005705767,"about_ca_system_score_gemma":0.003288122,"threshold_uncertainty_score":0.054163635},"labels":[],"label_agreement":null},{"id":"W2293317743","doi":"10.3233/978-1-61499-512-8-414","title":"An Ontology for Healthcare Quality Indicators: Challenges for Semantic Interoperability","year":2015,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Interoperability; Semantic interoperability; Computer science; Ontology; Quality (philosophy); Health care; Coding (social sciences); Set (abstract data type); Semantic heterogeneity; Data science; Knowledge management; World Wide Web; Semantic Web; Ontology-based data integration","score_opus":0.18870936671560049,"score_gpt":0.4696895068750612,"score_spread":0.2809801401594607,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293317743","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013047648,0.004148389,0.88368297,0.0830219,0.0006241555,0.0005555945,0.0019577811,0.00093211816,0.012029397],"genre_scores_gemma":[0.08299962,0.0024341566,0.9056246,0.0037511236,0.00027220382,0.0005731326,0.0028176403,0.00027749687,0.0012500127],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95940703,0.016389469,0.008734307,0.0029118594,0.011667938,0.0008893102],"domain_scores_gemma":[0.9147661,0.04939025,0.005894207,0.011566637,0.016038813,0.002344093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.051105518,0.00088170904,0.0021985373,0.009066832,0.004848846,0.013017911,0.0034799138,0.0037469245,0.0016983714],"category_scores_gemma":[0.06993176,0.00081970904,0.0030767603,0.015856713,0.007892847,0.03194281,0.005286261,0.0073131104,0.00068220217],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044620672,0.00013887195,0.0025733085,0.0007895074,0.0001105985,0.00020680063,0.0029953464,0.0041549336,0.0016592575,0.8713391,0.0137557015,0.102231964],"study_design_scores_gemma":[0.000043375938,0.000047180638,0.002157388,0.0019246432,0.00013659985,0.0004340124,0.004426766,0.03633406,0.0015341104,0.77677137,0.17603989,0.00015051627],"about_ca_topic_score_codex":0.018481873,"about_ca_topic_score_gemma":0.010391178,"teacher_disagreement_score":0.051105518,"about_ca_system_score_codex":0.009430377,"about_ca_system_score_gemma":0.017678034,"threshold_uncertainty_score":0.270275},"labels":[],"label_agreement":null},{"id":"W2293627165","doi":"","title":"Bridging Layperson's Queries with Medical Concepts- GRIUM@CLEF2015 eHealth Task 2","year":2015,"lang":"en","type":"article","venue":"CLEF (Working Notes)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Layperson; Computer science; Information retrieval; Task (project management); eHealth; Unified Medical Language System; Bridge (graph theory); Bridging (networking); Readability; Automatic summarization; World Wide Web; Data science; Health care; Medicine","score_opus":0.03709580862023832,"score_gpt":0.3136849472681744,"score_spread":0.2765891386479361,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2293627165","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6042812,0.0036438259,0.19254103,0.016605774,0.0016077958,0.009610724,0.093943335,0.022550454,0.055215895],"genre_scores_gemma":[0.53088224,0.0005998539,0.33550996,0.0027092926,0.00030002065,0.0044223564,0.10151339,0.0015719149,0.022490997],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99363786,0.0033711842,0.0005568152,0.0008190224,0.0011584602,0.00045656835],"domain_scores_gemma":[0.96216923,0.030372556,0.0008175434,0.0022918712,0.0030603139,0.001288423],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0086906925,0.00093649956,0.00096484635,0.002308411,0.0016330275,0.0020692083,0.000961392,0.002561318,0.018146016],"category_scores_gemma":[0.03295808,0.00029497652,0.0008735941,0.0013042551,0.00060972816,0.0018827395,0.0032258893,0.0011122222,0.005866965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036997588,0.0025320915,0.029766282,0.0042345873,0.00025206362,0.0049890685,0.008227745,0.007444594,0.04392832,0.009633992,0.4688889,0.41640246],"study_design_scores_gemma":[0.0022348415,0.0013711656,0.067792214,0.0007820956,0.00020952299,0.0046553183,0.014862082,0.10924333,0.16590434,0.01486749,0.61750704,0.0005706467],"about_ca_topic_score_codex":0.01230261,"about_ca_topic_score_gemma":0.012195213,"teacher_disagreement_score":0.018146016,"about_ca_system_score_codex":0.0018510725,"about_ca_system_score_gemma":0.0027550994,"threshold_uncertainty_score":0.06070447},"labels":[],"label_agreement":null},{"id":"W2296273445","doi":"","title":"Quality of care metric reporting from clinical narratives: Assessing ontology components","year":2014,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Metric (unit); Ontology; Health care; Quality (philosophy); Computer science; Data extraction; Data science; Quality management; Data quality; Knowledge management; Information retrieval; Data mining; MEDLINE; Business; Operations management; Engineering; Political science; Management system; Marketing","score_opus":0.13366557324830178,"score_gpt":0.4496565336896625,"score_spread":0.31599096044136077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296273445","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4680042,0.0022605476,0.48625666,0.006046216,0.00025263627,0.004817231,0.017371058,0.0015475996,0.01344388],"genre_scores_gemma":[0.5235136,0.00076232065,0.46216536,0.00029376824,0.000046098106,0.0014176856,0.011069332,0.00010401018,0.0006278266],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97332823,0.01144134,0.0058171507,0.0015590219,0.007412108,0.00044213227],"domain_scores_gemma":[0.80314267,0.14335857,0.02280266,0.0066374843,0.023113426,0.0009452121],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03073283,0.0005318908,0.00046363883,0.01149988,0.0011128581,0.004409979,0.0010504974,0.0007311566,0.0010105729],"category_scores_gemma":[0.1498461,0.00023404162,0.0009841603,0.009544805,0.0010155553,0.003413301,0.0030102648,0.0010427535,0.00025409655],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087498804,0.00053293054,0.30703792,0.0058254125,0.0008965607,0.0010798151,0.019323103,0.0146159055,0.015802752,0.035808574,0.013739271,0.5844628],"study_design_scores_gemma":[0.0002126639,0.0006869826,0.2760194,0.006723008,0.0016869442,0.0024803544,0.037286457,0.3610723,0.08787021,0.08382884,0.14160354,0.0005292672],"about_ca_topic_score_codex":0.006724521,"about_ca_topic_score_gemma":0.0063005746,"teacher_disagreement_score":0.9692672,"about_ca_system_score_codex":0.0032034286,"about_ca_system_score_gemma":0.005327814,"threshold_uncertainty_score":0.16253269},"labels":[],"label_agreement":null},{"id":"W2296712682","doi":"","title":"Mining the Biomedical Literature","year":2015,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Data science","score_opus":0.11449833389880208,"score_gpt":0.32301461238680773,"score_spread":0.20851627848800564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2296712682","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10629551,0.19121952,0.34299645,0.043572597,0.006836609,0.0029050019,0.11008429,0.005725899,0.1903641],"genre_scores_gemma":[0.18574835,0.11343504,0.5169123,0.005600847,0.004019268,0.0013788552,0.12756693,0.0007246012,0.044613685],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964259,0.00072533527,0.0005135628,0.00067689165,0.0014975786,0.00016056585],"domain_scores_gemma":[0.9915029,0.004319299,0.00079841685,0.0007687381,0.0023124327,0.00029828696],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0031152423,0.001020705,0.0011277765,0.032184765,0.0017796973,0.005881833,0.0019812558,0.0013979911,0.010120451],"category_scores_gemma":[0.021171976,0.00045641078,0.0014627617,0.022145456,0.0012418595,0.006253779,0.003154084,0.0013261699,0.0076572644],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014916644,0.00017468838,0.0063864505,0.006017095,0.00033269144,0.0026426874,0.0021540597,0.0019897416,0.006723362,0.03377084,0.12885073,0.8108084],"study_design_scores_gemma":[0.000047974117,0.000115232724,0.007806832,0.0035715965,0.0004264797,0.0037410809,0.0045424323,0.009381399,0.0068473876,0.10503158,0.85841316,0.00007478633],"about_ca_topic_score_codex":0.0022166101,"about_ca_topic_score_gemma":0.004691171,"teacher_disagreement_score":0.9678152,"about_ca_system_score_codex":0.0011395693,"about_ca_system_score_gemma":0.004248176,"threshold_uncertainty_score":0.033856273},"labels":[],"label_agreement":null},{"id":"W2297017727","doi":"","title":"A portal for the Canadian Virtual Health Library / Bibliothèque virtuelle canadienne de la santé (CVHL/BVCS): Accomplishments and possibilities, scope and structure - a review of best practices","year":2012,"lang":"en","type":"review","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Scope (computer science); Library science; Computer science","score_opus":0.017792052676392694,"score_gpt":0.2700910838545741,"score_spread":0.2522990311781814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2297017727","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008357351,0.65096813,0.013989631,0.0887033,0.003200947,0.0013529176,0.008863015,0.0020527656,0.22251192],"genre_scores_gemma":[0.07742917,0.78949195,0.06575502,0.011574339,0.0004946569,0.00063996576,0.009969305,0.00066829147,0.043977242],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98839206,0.0022954326,0.00086437515,0.00068847294,0.006766466,0.0009931222],"domain_scores_gemma":[0.97090125,0.007714256,0.0016508986,0.0011111269,0.015545569,0.0030769],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.016428933,0.0007919921,0.0008734572,0.016462445,0.0056673572,0.011892812,0.0026161713,0.002307202,0.012071526],"category_scores_gemma":[0.023111815,0.0005657333,0.0009839076,0.03180021,0.004661316,0.0061540157,0.0048595266,0.002013356,0.0038442363],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000082624014,0.00006581188,0.0022133254,0.014097863,0.00008125582,0.0002581583,0.004299218,0.0003490679,0.00072736497,0.039842915,0.20228794,0.73569447],"study_design_scores_gemma":[0.000007911147,0.000018571502,0.0040671425,0.007993726,0.000048094244,0.00013459915,0.0022261627,0.00006879624,0.00034554754,0.0012650444,0.9837855,0.000038902024],"about_ca_topic_score_codex":0.7786024,"about_ca_topic_score_gemma":0.8629373,"teacher_disagreement_score":0.9881072,"about_ca_system_score_codex":0.04391953,"about_ca_system_score_gemma":0.17391978,"threshold_uncertainty_score":0.44540286},"labels":[],"label_agreement":null},{"id":"W2314054433","doi":"10.5858/arpa.2012-0223-le","title":"Intersection of Surgical Pathology and Molecular Diagnostics for Targeted Therapy: Recommendation for Synoptic Reporting","year":2012,"lang":"en","type":"letter","venue":"Archives of Pathology & Laboratory Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Medicine; Surgical pathology; Molecular pathology; Medical diagnosis; Companion diagnostic; Pathology; General surgery; Medical physics; Cancer; Internal medicine; Biology","score_opus":0.024166466291903742,"score_gpt":0.3051569662644601,"score_spread":0.2809904999725564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2314054433","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00019475311,0.0031010523,0.0010554618,0.5696753,0.42382643,0.00007990087,0.00009066163,0.00019069188,0.0017857342],"genre_scores_gemma":[0.0015878996,0.0046618944,0.0029887143,0.4341232,0.5520721,0.00014403863,0.00013672425,0.00011382295,0.0041715964],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9730862,0.0077388985,0.006463323,0.0022700983,0.0092008505,0.0012406775],"domain_scores_gemma":[0.8037322,0.06023239,0.016937016,0.0066597937,0.1001986,0.012239911],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029175991,0.0016923195,0.0022440667,0.004207587,0.0027011724,0.008947938,0.006456399,0.027947599,0.010019699],"category_scores_gemma":[0.15166663,0.002059942,0.00235224,0.0021569568,0.005380126,0.011367983,0.003079608,0.03230388,0.01175717],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022509994,0.00001499546,0.00015441507,0.00017776457,0.000010262466,0.0001605234,0.000033584536,0.000029362345,0.000068363384,0.00044024896,0.9918389,0.0070491205],"study_design_scores_gemma":[0.00006863957,0.000048539896,0.0009815189,0.0016280947,0.00004248658,0.0012704342,0.00032506124,0.00037495815,0.00025567692,0.0024257337,0.9924925,0.00008639583],"about_ca_topic_score_codex":0.002482338,"about_ca_topic_score_gemma":0.004960656,"teacher_disagreement_score":0.029175991,"about_ca_system_score_codex":0.00366601,"about_ca_system_score_gemma":0.006573052,"threshold_uncertainty_score":0.15429926},"labels":[],"label_agreement":null},{"id":"W2317966583","doi":"10.2196/medinform.5275","title":"A Querying Method over RDF-ized Health Level Seven v2.5 Messages Using Life Science Knowledge Resources","year":2016,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; RDF; World Wide Web; Information retrieval; Database; Semantic Web","score_opus":0.062498714088168636,"score_gpt":0.40719834055889786,"score_spread":0.3446996264707292,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2317966583","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011965051,0.0001586823,0.9295826,0.00091109256,0.000070704475,0.0009822585,0.0062045967,0.04645371,0.0036712287],"genre_scores_gemma":[0.13230237,0.00026259857,0.83346975,0.00090332306,0.000053996515,0.0009964689,0.023841003,0.0043999995,0.0037705516],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99556327,0.00087191915,0.00074559357,0.0010492473,0.0015775593,0.00019245439],"domain_scores_gemma":[0.99530715,0.0023542026,0.00027570903,0.0009609214,0.00095864217,0.000143253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006054135,0.001421665,0.00089398545,0.003073277,0.0009830536,0.0021937755,0.0017115535,0.001230363,0.005551099],"category_scores_gemma":[0.010156313,0.0006864691,0.0024062388,0.001960773,0.0008231618,0.004584871,0.0027393692,0.0009670601,0.0014498085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019033159,0.0008480118,0.015855402,0.0026107526,0.00077709137,0.002897347,0.0044430434,0.050196953,0.08830728,0.13790308,0.12649521,0.5677627],"study_design_scores_gemma":[0.00050942256,0.00025091736,0.003918386,0.00030787487,0.00029561555,0.0010237364,0.0014317676,0.64383394,0.11559554,0.0519724,0.18054663,0.0003138204],"about_ca_topic_score_codex":0.010378818,"about_ca_topic_score_gemma":0.0071182232,"teacher_disagreement_score":0.010378818,"about_ca_system_score_codex":0.0015303289,"about_ca_system_score_gemma":0.0026667425,"threshold_uncertainty_score":0.032017708},"labels":[],"label_agreement":null},{"id":"W2318488401","doi":"10.1093/bioinformatics/btw155","title":"MOLGENIS/connect: a system for semi-automatic integration of heterogeneous phenotype data with applications in biobanks","year":2016,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; McGill University Health Centre","funders":"Terveyden ja hyvinvoinnin laitos; Nederlandse Organisatie voor Wetenschappelijk Onderzoek; European Commission","keywords":"Computer science; Biobank; Documentation; Source code; Data integration; Data mining; Terminology; Open source; Information retrieval; Ontology; Categorical variable; Task (project management); Data source; Matching (statistics); Software; Programming language; Machine learning; Bioinformatics","score_opus":0.024251684551665875,"score_gpt":0.2718240070330186,"score_spread":0.24757232248135275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2318488401","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02013474,0.00066386664,0.35870367,0.0010192548,0.00025335982,0.0011304755,0.04997556,0.5613034,0.0068157907],"genre_scores_gemma":[0.111026205,0.00060412847,0.68876797,0.0019024645,0.00022529275,0.0035245495,0.14354391,0.044225667,0.0061797304],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9968316,0.0007436209,0.00042234952,0.0011336361,0.00074261177,0.0001262927],"domain_scores_gemma":[0.99176604,0.004494541,0.00093248446,0.001666647,0.0007109348,0.00042943546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006365995,0.002132892,0.0012450407,0.005397011,0.0009174209,0.0029119859,0.0028511048,0.0012589978,0.025294712],"category_scores_gemma":[0.018640399,0.0014801438,0.0016954239,0.003537375,0.0007676438,0.003071574,0.005656826,0.0012843044,0.010033202],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004415947,0.0004507673,0.02808722,0.0033245673,0.0011785389,0.0022661649,0.0018579023,0.012007986,0.031893376,0.014013835,0.5380815,0.3624222],"study_design_scores_gemma":[0.0015912291,0.0004652497,0.033184957,0.0011015804,0.00052035874,0.003709823,0.0009076775,0.27273223,0.10026887,0.052475147,0.53250045,0.0005424732],"about_ca_topic_score_codex":0.002413324,"about_ca_topic_score_gemma":0.0027342986,"teacher_disagreement_score":0.025294712,"about_ca_system_score_codex":0.001352452,"about_ca_system_score_gemma":0.0022770737,"threshold_uncertainty_score":0.08461928},"labels":[],"label_agreement":null},{"id":"W2322383256","doi":"10.1177/1460458214555040","title":"Statistical classification of drug incidents due to look-alike sound-alike mix-ups","year":2014,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Japan Society for the Promotion of Science; Canadian Patient Safety Institute; University of Florida; Florida Agricultural and Mechanical University; Florida State University","keywords":"Support vector machine; Decision tree; Computer science; Feature selection; Feature (linguistics); Logistic regression; Receiver operating characteristic; Artificial intelligence; Statistical model; Data mining; Selection (genetic algorithm); Machine learning","score_opus":0.02653208338536603,"score_gpt":0.33263840672284933,"score_spread":0.30610632333748333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2322383256","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97880256,0.00035554118,0.015385836,0.0001766161,0.000043960332,0.00014328514,0.00328269,0.00050103856,0.0013084833],"genre_scores_gemma":[0.9825387,0.0001576873,0.0117454035,0.00004520744,0.000037996375,0.000054163527,0.0050100335,0.000019472724,0.000391279],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9984523,0.00021475439,0.00035578592,0.00025229817,0.0006138557,0.00011108898],"domain_scores_gemma":[0.9879274,0.0066373087,0.0028926942,0.0005326962,0.0017342507,0.00027562995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015237675,0.00039756048,0.0004451554,0.0080496,0.0002770807,0.0009164617,0.00036962627,0.00048614043,0.0006827485],"category_scores_gemma":[0.01247675,0.00009345626,0.0005348799,0.0039041277,0.00025159665,0.000856391,0.00055090425,0.0003332149,0.00037473312],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000632326,0.00028373578,0.7777744,0.00025123218,0.00019465199,0.0012574027,0.00049828517,0.009246045,0.011281771,0.00044255043,0.002458242,0.19567934],"study_design_scores_gemma":[0.00003033952,0.00043350944,0.7676065,0.00010618,0.00028207802,0.0020869416,0.0016670426,0.205152,0.016198618,0.0018604546,0.004507379,0.00006897414],"about_ca_topic_score_codex":0.0034745336,"about_ca_topic_score_gemma":0.0037502097,"teacher_disagreement_score":0.0080496,"about_ca_system_score_codex":0.0004562337,"about_ca_system_score_gemma":0.0005279706,"threshold_uncertainty_score":0.008058548},"labels":[],"label_agreement":null},{"id":"W2333066168","doi":"10.2196/resprot.5028","title":"Using Nonexperts for Annotating Pharmacokinetic Drug-Drug Interaction Mentions in Product Labeling: A Feasibility Study","year":2016,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; Fogarty International Center; National Institute on Aging; National Institutes of Health","keywords":"Drug; Product (mathematics); Pharmacokinetics; Computer science; Drug labeling; Pharmacology; Medicine","score_opus":0.5018194849802321,"score_gpt":0.6310378209124399,"score_spread":0.12921833593220777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2333066168","genre_codex":"empirical","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9760088,0.000064890504,0.0138949435,0.00030576997,0.00006412374,0.0069431765,0.00023814772,0.0002755713,0.002204586],"genre_scores_gemma":[0.8858073,0.00015408212,0.08966766,0.0011133596,0.0001229547,0.01916777,0.00077426434,0.00014106158,0.003051682],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9748199,0.01492246,0.0015652621,0.0044560283,0.002653875,0.0015824101],"domain_scores_gemma":[0.8861618,0.068504035,0.0047094943,0.011538102,0.021990748,0.0070958403],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04678256,0.0014919968,0.00090825476,0.0013617566,0.002462249,0.0022823634,0.002880674,0.0028464522,0.0070551983],"category_scores_gemma":[0.08011389,0.0013657712,0.0011141814,0.00049456174,0.0017570558,0.004633451,0.005832629,0.0020123157,0.0036107185],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.023087217,0.11643365,0.22133194,0.0053007244,0.00039335122,0.006294253,0.21866398,0.0035084703,0.08714854,0.0014521782,0.008784408,0.30760127],"study_design_scores_gemma":[0.018974083,0.18983774,0.41560277,0.001564828,0.0011960309,0.00911386,0.14751987,0.05411783,0.0716507,0.006990456,0.08182209,0.0016096856],"about_ca_topic_score_codex":0.0027301423,"about_ca_topic_score_gemma":0.005187161,"teacher_disagreement_score":0.95321745,"about_ca_system_score_codex":0.0010979631,"about_ca_system_score_gemma":0.0028748303,"threshold_uncertainty_score":0.2474128},"labels":[],"label_agreement":null},{"id":"W2335766810","doi":"10.29173/jchla/jabsc.v32i3.27566","title":"Canadian Virtual Health Library / Bibliotheque virtuelle canadienne de la sante update","year":2011,"lang":"fr","type":"article","venue":"Journal of the Canadian Health Libraries Association / Journal de l Association de bilbiothèques de la santé du Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Political science; Art; Computer science","score_opus":0.005439048733149942,"score_gpt":0.227397975561086,"score_spread":0.22195892682793608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2335766810","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022045514,0.22335438,0.0030456511,0.24167876,0.054872695,0.00037606456,0.046191383,0.002051047,0.42622545],"genre_scores_gemma":[0.040015433,0.47056183,0.019047868,0.048519444,0.020647695,0.00048935134,0.08265036,0.0010305039,0.3170375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99098116,0.0006839973,0.00066850433,0.00033482612,0.006589777,0.00074173603],"domain_scores_gemma":[0.9529514,0.004324796,0.0013546693,0.0011926714,0.03508732,0.0050892355],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.008583777,0.0010955407,0.001631016,0.02137561,0.004932527,0.011110418,0.003902995,0.0035608064,0.052943025],"category_scores_gemma":[0.032308355,0.00066581776,0.0013911779,0.031264547,0.0026278596,0.004104934,0.003232672,0.0041724946,0.01167737],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017213466,0.000011240267,0.00028574673,0.00041404515,0.000013010121,0.00003265539,0.00003621485,0.000072620074,0.00002518561,0.003348073,0.93966717,0.056076862],"study_design_scores_gemma":[0.000006831081,0.0000020997259,0.0011750102,0.00038410234,0.00001614965,0.00005837272,0.000046478024,0.000037474045,0.000031730582,0.0005259578,0.9977054,0.000010439316],"about_ca_topic_score_codex":0.8853558,"about_ca_topic_score_gemma":0.92141217,"teacher_disagreement_score":0.9888896,"about_ca_system_score_codex":0.056015957,"about_ca_system_score_gemma":0.19406825,"threshold_uncertainty_score":0.40642607},"labels":[],"label_agreement":null},{"id":"W2337837203","doi":"10.1101/034397","title":"Event Extraction from Biomedical Literature","year":2015,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"BC Cancer Agency; University of British Columbia; Genome British Columbia; Canada's Michael Smith Genome Sciences Centre; Genome Canada","keywords":"Biomedical text mining; Computer science; Scope (computer science); Event (particle physics); Task (project management); Information extraction; Data science; Field (mathematics); Domain (mathematical analysis); Information retrieval; Named-entity recognition; Natural language processing; Artificial intelligence; Text mining","score_opus":0.01498455383574694,"score_gpt":0.26044739010449697,"score_spread":0.24546283626875004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2337837203","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037505712,0.027158076,0.6687019,0.005932182,0.001636249,0.003430229,0.20838393,0.021965234,0.025286485],"genre_scores_gemma":[0.12494271,0.014736514,0.6065316,0.00089355797,0.0013692097,0.0020642867,0.24324112,0.000912988,0.0053080306],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9959359,0.000879516,0.00082692096,0.0010007935,0.0012061689,0.0001507159],"domain_scores_gemma":[0.9797309,0.013372565,0.0018567066,0.0016037837,0.0030446325,0.00039139812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036315136,0.0016523398,0.0011463615,0.025574327,0.0012798498,0.0035156256,0.0016890857,0.0014281892,0.008785921],"category_scores_gemma":[0.028924217,0.0005645076,0.0016903003,0.015781783,0.0006603631,0.0035305426,0.0031599742,0.0014027434,0.005365016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000559342,0.00021332536,0.011705168,0.016585037,0.00043442287,0.0046386984,0.0013163757,0.0059634303,0.019398399,0.027578002,0.11354245,0.7980653],"study_design_scores_gemma":[0.00025469033,0.00021826616,0.024017587,0.0037363064,0.00103799,0.005441724,0.0019512046,0.08595303,0.044649675,0.14647363,0.6860104,0.00025540713],"about_ca_topic_score_codex":0.0020201355,"about_ca_topic_score_gemma":0.001979664,"teacher_disagreement_score":0.025574327,"about_ca_system_score_codex":0.0010659737,"about_ca_system_score_gemma":0.002929285,"threshold_uncertainty_score":0.029391885},"labels":[],"label_agreement":null},{"id":"W2338265029","doi":"10.1186/s13326-016-0062-4","title":"VICO: Ontology-based representation and integrative analysis of Vaccination Informed Consent forms","year":2016,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Medical School, University of Michigan; University of Michigan","keywords":"Informed consent; Ontology; Vaccination; Standardization; Computer science; Knowledge management; Family medicine; Medicine; Medical education; Alternative medicine; Epistemology; Immunology; Pathology","score_opus":0.022244863642136604,"score_gpt":0.33351079087237123,"score_spread":0.3112659272302346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2338265029","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013322689,0.00046730053,0.9324988,0.00240269,0.00018470643,0.0018546763,0.015331849,0.01146426,0.02247302],"genre_scores_gemma":[0.0763186,0.0008637536,0.8723105,0.0005163388,0.00006311325,0.0012265426,0.04301524,0.0015261025,0.004159725],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99302626,0.0020763983,0.0014553415,0.000911172,0.0021748268,0.0003560184],"domain_scores_gemma":[0.9885494,0.004809323,0.0012929887,0.0027253376,0.0021100887,0.0005127731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00909089,0.0011740023,0.00069468905,0.009739396,0.0023946385,0.0072268704,0.0024159842,0.0015419341,0.0043745125],"category_scores_gemma":[0.023661574,0.0009347497,0.003408458,0.007574502,0.0019041498,0.008806125,0.0070005464,0.0026008796,0.0011493695],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026134704,0.0003933273,0.014109104,0.0019656213,0.00035548312,0.0020309326,0.01577171,0.03440465,0.009147005,0.607011,0.053960774,0.2605889],"study_design_scores_gemma":[0.00006540494,0.00008217305,0.0059524053,0.0015424744,0.00026112673,0.0010367987,0.006589278,0.16209048,0.0070661968,0.23812598,0.5769826,0.0002050448],"about_ca_topic_score_codex":0.04247771,"about_ca_topic_score_gemma":0.051567458,"teacher_disagreement_score":0.04247771,"about_ca_system_score_codex":0.0052428017,"about_ca_system_score_gemma":0.010682152,"threshold_uncertainty_score":0.084460914},"labels":[],"label_agreement":null},{"id":"W2338820613","doi":"10.1101/049205","title":"Omics Discovery Index - Discovering and Linking Public ‘Omics’ Datasets","year":2016,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"Biotechnology and Biological Sciences Research Council; National Institutes of Health; National Natural Science Foundation of China; Wellcome Trust","keywords":"Computer science; Metadata; Omics; Data science; Identification (biology); Interoperability; Data sharing; World Wide Web; Bioinformatics; Biology","score_opus":0.015929196029360805,"score_gpt":0.2351394664125911,"score_spread":0.2192102703832303,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2338820613","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014362929,0.015367444,0.4974348,0.0093298685,0.0019876203,0.0032049809,0.31827828,0.08941591,0.05061824],"genre_scores_gemma":[0.032217745,0.009798484,0.48955426,0.0029759866,0.00063374196,0.002383837,0.44866627,0.006658692,0.0071109748],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9828866,0.0032829142,0.002273195,0.003788232,0.0067957877,0.00097325153],"domain_scores_gemma":[0.9833673,0.005112188,0.0017723708,0.006146373,0.0024123099,0.001189492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018051138,0.0035648507,0.0036183095,0.03179077,0.0044775186,0.010508807,0.006125733,0.0029587734,0.014244605],"category_scores_gemma":[0.042588446,0.0021074275,0.0054115024,0.026085181,0.0017584865,0.013542402,0.026813244,0.0049629514,0.020740397],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011154339,0.00042766434,0.02376788,0.009206905,0.0025490045,0.0015535834,0.0022720187,0.004578956,0.016006926,0.15372732,0.43341035,0.35138392],"study_design_scores_gemma":[0.00013125296,0.00013167536,0.009110559,0.0026176397,0.00062262686,0.0010911779,0.00081485906,0.015137595,0.014677691,0.13348585,0.82191974,0.00025938568],"about_ca_topic_score_codex":0.00638823,"about_ca_topic_score_gemma":0.008563656,"teacher_disagreement_score":0.03179077,"about_ca_system_score_codex":0.0032198967,"about_ca_system_score_gemma":0.01186395,"threshold_uncertainty_score":0.09546465},"labels":[],"label_agreement":null},{"id":"W2339286586","doi":"10.1177/1049732315611266","title":"Knowledge Translation","year":2015,"lang":"en","type":"editorial","venue":"Qualitative Health Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia, Okanagan Campus; University of British Columbia","funders":"","keywords":"Psychology; Medicine; Computer science","score_opus":0.6504237998405149,"score_gpt":0.6549829815038505,"score_spread":0.004559181663335576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2339286586","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000086717686,0.014647885,0.00379971,0.24074672,0.7156262,0.00016362872,0.00073520135,0.00035093114,0.023843087],"genre_scores_gemma":[0.004878712,0.02880637,0.0065549957,0.26677814,0.5471867,0.00077239407,0.0018175037,0.0007996143,0.14240566],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9814454,0.007947402,0.0030812777,0.0012595861,0.0054511894,0.00081527117],"domain_scores_gemma":[0.93064576,0.03959696,0.0027114009,0.0053759874,0.01914557,0.0025243599],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019342462,0.0014320981,0.0018877983,0.007183027,0.0026586624,0.0120507525,0.0028338092,0.011644908,0.08440628],"category_scores_gemma":[0.094723955,0.0009079765,0.0026510656,0.0033218137,0.004376031,0.0068152994,0.0073729237,0.014704456,0.03146211],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002643129,0.000007301306,0.000007948579,0.00082825904,0.000024546098,0.00007393233,0.00012908127,0.000024019919,0.00007528389,0.0065600965,0.9643266,0.02791648],"study_design_scores_gemma":[0.000027271853,0.00000500148,0.000033635835,0.0011429063,0.00001917768,0.000063656706,0.00007214151,0.000041117815,0.00008997905,0.007224713,0.9912718,0.000008487406],"about_ca_topic_score_codex":0.003200896,"about_ca_topic_score_gemma":0.0048186486,"teacher_disagreement_score":0.08440628,"about_ca_system_score_codex":0.005192431,"about_ca_system_score_gemma":0.0102610355,"threshold_uncertainty_score":0.28236717},"labels":[],"label_agreement":null},{"id":"W2341413855","doi":"10.22374/cjgim.v8i3.65","title":"Lessons from History: Still Relevant in the “Information Age”","year":2013,"lang":"en","type":"article","venue":"Canadian Journal of General Internal Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine; Specialty; Disease; Medical history; Physical examination; Epistemology; Work (physics); Classics; Family medicine; History; Surgery; Pathology; Engineering; Philosophy","score_opus":0.027746380989015537,"score_gpt":0.2649083223426648,"score_spread":0.23716194135364926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2341413855","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00045131682,0.04971855,0.0024560648,0.9242652,0.01789169,0.000009388446,0.000121508565,0.000051325827,0.005034851],"genre_scores_gemma":[0.03807764,0.21700199,0.011416076,0.5880383,0.13383529,0.00008555686,0.0006464477,0.0003901173,0.0105085345],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99249446,0.0037911113,0.0008052398,0.00048297364,0.0020224527,0.00040377592],"domain_scores_gemma":[0.9090709,0.07828227,0.0020018544,0.0024215523,0.004694637,0.0035288306],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017476775,0.00068123697,0.00113769,0.005883293,0.004964918,0.021564657,0.0028823153,0.008510317,0.0051602153],"category_scores_gemma":[0.074160755,0.0006874393,0.0008558444,0.0062081753,0.01896658,0.045000765,0.006626994,0.0155888125,0.001844532],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034243087,0.00001690295,0.0007807851,0.0009834449,0.000044144283,0.00032781382,0.008183253,0.00008154785,0.00008654264,0.080261834,0.75250113,0.1566983],"study_design_scores_gemma":[0.0000072025637,0.000008141494,0.00045490754,0.001460021,0.000013386059,0.0004322107,0.0043737725,0.00007110371,0.000053008516,0.09979155,0.8933046,0.000030150244],"about_ca_topic_score_codex":0.0062873513,"about_ca_topic_score_gemma":0.021337202,"teacher_disagreement_score":0.021564657,"about_ca_system_score_codex":0.0052817524,"about_ca_system_score_gemma":0.009680053,"threshold_uncertainty_score":0.092427135},"labels":[],"label_agreement":null},{"id":"W2344996561","doi":"10.1186/s13104-016-2023-5","title":"Integrating text mining, data mining, and network analysis for identifying genetic breast cancer trends","year":2016,"lang":"en","type":"article","venue":"BMC Research Notes","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Breast cancer; Data science; Disease; Computer science; Computational biology; Cancer; Bioinformatics; Data mining; Gene; Medicine; Biology; Genetics; Pathology","score_opus":0.22262810182145837,"score_gpt":0.45631828068217417,"score_spread":0.2336901788607158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2344996561","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48011544,0.010198494,0.445847,0.0057757953,0.0002955622,0.002139524,0.039538965,0.006108635,0.009980579],"genre_scores_gemma":[0.4877686,0.0034785504,0.4841655,0.00031414995,0.0001947551,0.00086492347,0.022014886,0.00011351854,0.0010851417],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99832183,0.0005761992,0.00033815554,0.00028405437,0.0003987183,0.00008116142],"domain_scores_gemma":[0.9902052,0.0072527113,0.0011570571,0.0002848519,0.0009015021,0.00019866913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033358496,0.0008887644,0.00074149173,0.018613229,0.0006588083,0.0018785122,0.0006782826,0.0005197293,0.0010571148],"category_scores_gemma":[0.008996016,0.0001910624,0.0011766668,0.013138777,0.00024273778,0.0022741768,0.00087072083,0.0006387259,0.00043858352],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044481567,0.0009819755,0.21094856,0.0033123062,0.0014751175,0.0014549047,0.0015714725,0.025550567,0.010634275,0.006024147,0.009443457,0.72815835],"study_design_scores_gemma":[0.00012171911,0.0005971135,0.18303578,0.0013448591,0.001823255,0.001907507,0.004548129,0.70047367,0.01712688,0.049917392,0.038903102,0.00020063936],"about_ca_topic_score_codex":0.0055601057,"about_ca_topic_score_gemma":0.009235408,"teacher_disagreement_score":0.018613229,"about_ca_system_score_codex":0.0009038738,"about_ca_system_score_gemma":0.00130933,"threshold_uncertainty_score":0.017641842},"labels":[],"label_agreement":null},{"id":"W2347198336","doi":"10.5281/zenodo.50971","title":"corset_pipeline: First complete release","year":2016,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Pipeline (software); Computer science; Programming language","score_opus":0.03560564919271993,"score_gpt":0.24871502057675254,"score_spread":0.2131093713840326,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2347198336","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011050304,0.00043945658,0.053056877,0.00026471665,0.0004209608,0.00032682225,0.398817,0.53734326,0.008225793],"genre_scores_gemma":[0.0055682673,0.0002534551,0.054742914,0.00066552195,0.00015163014,0.0013822477,0.7388159,0.19003119,0.008388879],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99783164,0.0002681155,0.00022573993,0.00092357263,0.0005240491,0.00022682885],"domain_scores_gemma":[0.9970878,0.0009970044,0.00019214803,0.0007617088,0.0007735115,0.000187797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003147665,0.005388814,0.0026637437,0.0036244327,0.0018885613,0.0036859242,0.0037374082,0.00174829,0.15570347],"category_scores_gemma":[0.01060397,0.003158168,0.0037399463,0.0028602837,0.0006949519,0.0034773585,0.004396062,0.003598465,0.16771229],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004459646,0.000040896248,0.0009287974,0.0012631259,0.0002031409,0.00007678292,0.00015401132,0.000474048,0.004930215,0.0013592471,0.9719982,0.018125514],"study_design_scores_gemma":[0.00041762472,0.00006678791,0.0027512258,0.00029460088,0.00017994716,0.0002306611,0.000062164836,0.004790478,0.01723893,0.008210159,0.965524,0.00023345865],"about_ca_topic_score_codex":0.006845486,"about_ca_topic_score_gemma":0.0075597665,"teacher_disagreement_score":0.15570347,"about_ca_system_score_codex":0.0012597098,"about_ca_system_score_gemma":0.0033012095,"threshold_uncertainty_score":0.52088},"labels":[],"label_agreement":null},{"id":"W2352364794","doi":"","title":"Curriculum Mapping in Canadian and UK Medical Schools","year":2010,"lang":"en","type":"article","venue":"Fudan jiaoyu luntan","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Curriculum; Construct (python library); Process (computing); Medical education; Upgrade; Interface (matter); Political science; Medicine; Sociology; Computer science; Pedagogy","score_opus":0.006442747712204863,"score_gpt":0.2507160061811567,"score_spread":0.24427325846895181,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2352364794","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8540054,0.0068680365,0.0017680933,0.024049109,0.00043708115,0.00039771653,0.0047287405,0.0004033256,0.107342586],"genre_scores_gemma":[0.9728696,0.0036602765,0.002888998,0.0014600983,0.000035017983,0.00009386379,0.0016510476,0.000072853014,0.017268239],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9891033,0.0015344057,0.0006210546,0.0011452413,0.0048482376,0.0027477408],"domain_scores_gemma":[0.96777827,0.0029054505,0.0026773182,0.0006401334,0.015638517,0.0103602605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006494361,0.00024079319,0.0003647811,0.007700739,0.007207621,0.0046317247,0.0013153928,0.00067544845,0.0057178554],"category_scores_gemma":[0.025890794,0.00037518225,0.00035211787,0.020545082,0.0017791985,0.0015612007,0.0036599005,0.00092255627,0.0005535708],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028719698,0.00019526032,0.32133573,0.0010197366,0.000042019572,0.0003914965,0.092977025,0.00078759016,0.0015845994,0.021919347,0.06767535,0.4917846],"study_design_scores_gemma":[0.000019122675,0.000094388844,0.6811152,0.00036123913,0.000025897074,0.00020295051,0.05509213,0.00043574674,0.0007095588,0.00051840325,0.26134953,0.00007591767],"about_ca_topic_score_codex":0.964315,"about_ca_topic_score_gemma":0.9802698,"teacher_disagreement_score":0.9244626,"about_ca_system_score_codex":0.07553736,"about_ca_system_score_gemma":0.16385654,"threshold_uncertainty_score":0.5480645},"labels":[],"label_agreement":null},{"id":"W2356882517","doi":"10.1093/jamia/ocw028","title":"Learning statistical models of phenotypes using noisy labeled training data","year":2016,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":165,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"U.S. National Library of Medicine; National Institute of General Medical Sciences; National Human Genome Research Institute","keywords":"Computer science; Scalability; Machine learning; Artificial intelligence; Logistic regression; Feature (linguistics); Phenotype; Feature engineering; Implementation; Predictive modelling; Data mining; Deep learning; Database","score_opus":0.04244371233329072,"score_gpt":0.31999571476336974,"score_spread":0.27755200243007905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2356882517","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.101875655,0.00018483626,0.8935858,0.000583356,0.000029435025,0.0001464031,0.0014603428,0.001462473,0.0006716987],"genre_scores_gemma":[0.67323565,0.00016916892,0.31737295,0.00040808553,0.00006600975,0.0006786002,0.0070422185,0.00015924372,0.0008681259],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994936,0.0027078516,0.00032585624,0.0013210678,0.0005832041,0.00012611819],"domain_scores_gemma":[0.9611462,0.029558424,0.0032042966,0.003392693,0.002368602,0.00032989413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009586131,0.0013053176,0.0010123999,0.0018734024,0.00049926824,0.0018918872,0.0017222614,0.0011564454,0.0008825135],"category_scores_gemma":[0.04012385,0.0005809961,0.0011338784,0.0012296828,0.0011064075,0.0017439075,0.0011991485,0.0017337837,0.0005646524],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027446484,0.00033083803,0.043759216,0.00018642572,0.0002939466,0.000339393,0.0004146907,0.8568134,0.0019685498,0.0062000584,0.0033079688,0.08611118],"study_design_scores_gemma":[0.000024694828,0.000049973965,0.0022650815,0.000036034155,0.000027294747,0.000045004435,0.000042928463,0.9855378,0.00089625694,0.010549677,0.0005133454,0.000011918496],"about_ca_topic_score_codex":0.0038886652,"about_ca_topic_score_gemma":0.006047883,"teacher_disagreement_score":0.009586131,"about_ca_system_score_codex":0.0014320879,"about_ca_system_score_gemma":0.0017196283,"threshold_uncertainty_score":0.05069691},"labels":[],"label_agreement":null},{"id":"W2365672368","doi":"","title":"An Investigation of the Eectiveness of Concept-based Approach in Medical Information Retrieval GRIUM @ CLEF2014eHealthTask 3","year":2014,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Clef; Computer science; Information retrieval; Task (project management); Unified Medical Language System; Natural language processing; Artificial intelligence; Resource (disambiguation); Domain (mathematical analysis)","score_opus":0.011150939771294941,"score_gpt":0.26254883015936537,"score_spread":0.2513978903880704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2365672368","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7532152,0.0283579,0.17613988,0.0044017076,0.00084245496,0.003107715,0.0029376175,0.0046683922,0.026329208],"genre_scores_gemma":[0.7016866,0.004081971,0.28532282,0.0008510703,0.00035168373,0.00080483634,0.0023951742,0.0003232915,0.004182554],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9666577,0.02207341,0.002339287,0.0020257414,0.0062767263,0.0006271771],"domain_scores_gemma":[0.8105449,0.17451645,0.0029194003,0.004104927,0.006764886,0.0011494068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.044136517,0.0020042649,0.0014232744,0.0075510466,0.001101779,0.0035744996,0.0025273846,0.0029693465,0.0027784281],"category_scores_gemma":[0.094727986,0.00059990794,0.0011638569,0.004603745,0.0013940452,0.006313072,0.003511018,0.0021021822,0.0010314225],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0068011414,0.0044287937,0.031176122,0.0059832325,0.0016635408,0.00045502713,0.002955934,0.022779217,0.026728773,0.0048581557,0.0110408785,0.88112926],"study_design_scores_gemma":[0.0024699583,0.014918232,0.0902259,0.00091938494,0.0016294157,0.0020211216,0.004825125,0.77679837,0.06843419,0.012488797,0.02463036,0.00063916296],"about_ca_topic_score_codex":0.0070931097,"about_ca_topic_score_gemma":0.0045970636,"teacher_disagreement_score":0.044136517,"about_ca_system_score_codex":0.0017766064,"about_ca_system_score_gemma":0.0014170125,"threshold_uncertainty_score":0.23341894},"labels":[],"label_agreement":null},{"id":"W2396211122","doi":"10.1007/978-1-60761-854-6_20","title":"Annotating the Regulatory Genome","year":2010,"lang":"en","type":"review","venue":"Methods in molecular biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"","keywords":"Biology; Genome; Function (biology); Context (archaeology); Gene; Computational biology; Regulatory sequence; Genetics; Regulation of gene expression","score_opus":0.054161220254314066,"score_gpt":0.46008746885436047,"score_spread":0.4059262486000464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396211122","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0115870405,0.7681598,0.15202117,0.0056694667,0.0016907775,0.00040433856,0.020373957,0.002597382,0.037496053],"genre_scores_gemma":[0.022601742,0.6638782,0.21751986,0.0021209796,0.00044212534,0.0003991581,0.07556281,0.00056464894,0.016910456],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994081,0.000108831846,0.0000700168,0.00018989609,0.00018375492,0.00003944313],"domain_scores_gemma":[0.9986124,0.00055559864,0.00021262543,0.00020015059,0.00037957722,0.00003959882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017640062,0.0014918305,0.00212029,0.0049513043,0.00059540104,0.0017405784,0.002251705,0.0013606527,0.0053875945],"category_scores_gemma":[0.0029022372,0.00046242127,0.0011221466,0.007683956,0.0007592461,0.0023241288,0.0010988708,0.0013333294,0.0057434677],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000086207496,0.000044778913,0.001151574,0.01411907,0.00012396133,0.0005646321,0.00036243277,0.0018712811,0.027612627,0.020704614,0.038886253,0.89447266],"study_design_scores_gemma":[0.000009478828,0.000016257158,0.0024073818,0.0020867232,0.00015533007,0.0011500802,0.000119870194,0.00091694976,0.008383551,0.00899829,0.9757284,0.00002767718],"about_ca_topic_score_codex":0.005938842,"about_ca_topic_score_gemma":0.0056729843,"teacher_disagreement_score":0.005938842,"about_ca_system_score_codex":0.0014833474,"about_ca_system_score_gemma":0.0029576076,"threshold_uncertainty_score":0.018023312},"labels":[],"label_agreement":null},{"id":"W2396432520","doi":"","title":"BIM: an open ontology for the annotation of biomedical images.","year":2015,"lang":"en","type":"article","venue":"ICBO","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Ontology; Computer science; Automatic image annotation; Annotation; Information retrieval; Image retrieval; RDF; Controlled vocabulary; Representation (politics); Open Biomedical Ontologies; Image (mathematics); Semantic Web; World Wide Web; Artificial intelligence; Ontology-based data integration; Suggested Upper Merged Ontology","score_opus":0.08984978563287857,"score_gpt":0.3811510993648924,"score_spread":0.29130131373201384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396432520","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002842073,0.0010683236,0.8906821,0.0018807113,0.00037495335,0.0016327192,0.04160364,0.02971922,0.03019625],"genre_scores_gemma":[0.030901643,0.0022311092,0.8210702,0.0020845775,0.00018991792,0.002946359,0.12215052,0.004828182,0.013597547],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9951233,0.0009401563,0.00091899524,0.0005335008,0.0022107342,0.0002732139],"domain_scores_gemma":[0.9949142,0.001406347,0.0005718688,0.0012666114,0.0014005556,0.0004403811],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006293739,0.0010512002,0.00095330825,0.008744983,0.001993793,0.0044146213,0.0034998362,0.0022982212,0.006204496],"category_scores_gemma":[0.010419292,0.0009962608,0.0021680025,0.0062157786,0.0018402736,0.007683531,0.005644734,0.0027754083,0.006645962],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037328317,0.00045780293,0.003411694,0.0035903687,0.00029103487,0.0015323543,0.003344859,0.007545066,0.025246508,0.42025435,0.26496065,0.26899207],"study_design_scores_gemma":[0.000040978368,0.00003255529,0.0015928576,0.000771483,0.00006560751,0.0008542723,0.00048044982,0.013781827,0.0075523257,0.05878733,0.915947,0.00009337076],"about_ca_topic_score_codex":0.015536345,"about_ca_topic_score_gemma":0.01624919,"teacher_disagreement_score":0.015536345,"about_ca_system_score_codex":0.0031044655,"about_ca_system_score_gemma":0.006718905,"threshold_uncertainty_score":0.033284903},"labels":[],"label_agreement":null},{"id":"W2396634530","doi":"","title":"SEBI: An Architecture for Biomedical Image Discovery, Interoperability and Reusability Based on Semantic Enrichment.","year":2014,"lang":"en","type":"article","venue":"SWAT4LS","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Interoperability; Information retrieval; Reusability; Key (lock); Information extraction; World Wide Web; Annotation; Semantic search; Semantic Web; Artificial intelligence; Software","score_opus":0.010080140793564967,"score_gpt":0.27974026627342813,"score_spread":0.2696601254798632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2396634530","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002438434,0.0005744634,0.93699604,0.0005994603,0.000070440525,0.00043997017,0.0015571915,0.05399605,0.003327884],"genre_scores_gemma":[0.02662702,0.0010151969,0.9440066,0.00066945923,0.000053122396,0.00069919234,0.017145487,0.0030894727,0.006694433],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9977034,0.00041101073,0.0002972082,0.0005454646,0.0008600768,0.00018294985],"domain_scores_gemma":[0.99571747,0.0010237119,0.0004025984,0.0014506009,0.0009740332,0.00043157488],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.0057518417,0.001557321,0.0017506183,0.009643725,0.0014703023,0.005129296,0.0052526314,0.002341413,0.004101867],"category_scores_gemma":[0.009977587,0.0015248539,0.0028150063,0.007265952,0.0020843302,0.009877785,0.010218614,0.0030368157,0.0066355],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013270184,0.0009634114,0.006465118,0.0028560034,0.000983551,0.0017052741,0.0035576315,0.012740921,0.051814266,0.118100345,0.12919155,0.6702949],"study_design_scores_gemma":[0.00011344567,0.00025092054,0.0050870664,0.0006636329,0.0004332127,0.001992214,0.0009581154,0.26205605,0.060101323,0.20180602,0.46620846,0.0003295787],"about_ca_topic_score_codex":0.00744838,"about_ca_topic_score_gemma":0.010724408,"teacher_disagreement_score":0.99474734,"about_ca_system_score_codex":0.0018115521,"about_ca_system_score_gemma":0.002596318,"threshold_uncertainty_score":0.030419052},"labels":[],"label_agreement":null},{"id":"W2397149926","doi":"10.1007/978-1-4939-0709-0_5","title":"Biological Information Extraction and Co-occurrence Analysis","year":2014,"lang":"en","type":"review","venue":"Methods in molecular biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":43,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Inference; Data science; Field (mathematics); Representation (politics); String (physics); Information extraction; Information retrieval; Data mining; Artificial intelligence; Mathematics","score_opus":0.06840891714613861,"score_gpt":0.5014912274252321,"score_spread":0.4330823102790935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397149926","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0055820253,0.466728,0.4985738,0.0030205813,0.0013161756,0.00065558375,0.008662561,0.0038423813,0.011618815],"genre_scores_gemma":[0.029711323,0.4632049,0.47367126,0.0014492534,0.0013544834,0.0010525424,0.01983966,0.00052639865,0.009190165],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977449,0.00032077284,0.0003723444,0.00049563585,0.00096722675,0.00009908416],"domain_scores_gemma":[0.9961552,0.0023880946,0.00037425544,0.00031983445,0.00068767736,0.00007488437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002621365,0.0019785697,0.0032291322,0.015580641,0.0005236891,0.0024839053,0.0027467704,0.0010762054,0.0027247327],"category_scores_gemma":[0.005320208,0.00056137575,0.0023605318,0.018772291,0.0010557345,0.003019281,0.002255718,0.0018807182,0.0036108682],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004805648,0.0000746651,0.00068064925,0.005012474,0.0001904568,0.00011779583,0.00009920151,0.0006282355,0.0032692044,0.0061628,0.010481037,0.9732354],"study_design_scores_gemma":[0.00007383924,0.00012208238,0.015073104,0.0072948798,0.0016193336,0.0048107905,0.0005501008,0.02437868,0.032978028,0.12507144,0.78771603,0.00031171794],"about_ca_topic_score_codex":0.001839982,"about_ca_topic_score_gemma":0.0015890797,"teacher_disagreement_score":0.015580641,"about_ca_system_score_codex":0.0008026999,"about_ca_system_score_gemma":0.0020572464,"threshold_uncertainty_score":0.0138632655},"labels":[],"label_agreement":null},{"id":"W2397593430","doi":"","title":"A landmark in biomedical information: many ways are leading to PubMed - MediaWiki tags open remote literature access to PubMed.","year":2011,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Landmark; World Wide Web; Computer science; Information retrieval; Data science; Internet privacy; Medicine; Artificial intelligence","score_opus":0.044184510651065795,"score_gpt":0.2744967147514106,"score_spread":0.23031220410034478,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2397593430","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010555529,0.03560414,0.15568481,0.46470627,0.075621106,0.00079691556,0.032438174,0.04125291,0.18334015],"genre_scores_gemma":[0.09135841,0.0339433,0.4004354,0.11578329,0.03779562,0.0013118623,0.054569438,0.019628534,0.24517424],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9893355,0.0034789233,0.002401422,0.0008908307,0.0032919508,0.0006013555],"domain_scores_gemma":[0.90796405,0.03818103,0.012972016,0.015930787,0.015735649,0.009216626],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.017683277,0.0010081499,0.0012903117,0.021519445,0.0036868965,0.013268307,0.0024099995,0.0052144537,0.07855049],"category_scores_gemma":[0.08339932,0.0010722971,0.0012279885,0.021991178,0.0039054893,0.028397445,0.01117252,0.007289498,0.057260215],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004218663,0.00006953367,0.0015925063,0.0037592398,0.00014907964,0.0010631247,0.001713972,0.00007299809,0.0060509783,0.06892599,0.7360741,0.18010668],"study_design_scores_gemma":[0.000038078913,0.000035147354,0.00090296386,0.0008675351,0.000053911343,0.0004146759,0.0005762445,0.00007964237,0.0020268576,0.011991147,0.9829571,0.0000566928],"about_ca_topic_score_codex":0.003092653,"about_ca_topic_score_gemma":0.0066132816,"teacher_disagreement_score":0.9867317,"about_ca_system_score_codex":0.002199354,"about_ca_system_score_gemma":0.009658004,"threshold_uncertainty_score":0.26277757},"labels":[],"label_agreement":null},{"id":"W2398249934","doi":"10.3233/978-1-61499-289-9-1024","title":"Machine Learning Methods for Clinical Forms Analysis in Mental Health","year":2013,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Centre for Addiction and Mental Health","funders":"","keywords":"Computer science; Support vector machine; Artificial intelligence; Machine learning; Task (project management); Process (computing); Transformation (genetics); Feature (linguistics); Data mining; Natural language processing","score_opus":0.07718247409998508,"score_gpt":0.4980112670103522,"score_spread":0.4208287929103671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398249934","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031947773,0.0019164424,0.990203,0.00089854875,0.000078633646,0.00022392866,0.00033987217,0.0016337516,0.0015110034],"genre_scores_gemma":[0.07443059,0.0013899662,0.92022914,0.00024069828,0.00024686096,0.0006795779,0.0009181124,0.00017423766,0.0016908451],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924185,0.004395694,0.0007766382,0.00090794085,0.0013345759,0.0001665574],"domain_scores_gemma":[0.9797332,0.015712915,0.0010796265,0.0012683665,0.002014943,0.00019100754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010534617,0.0010277207,0.0014933058,0.006804947,0.00082006096,0.003011557,0.0020054013,0.0015982449,0.003536607],"category_scores_gemma":[0.029913692,0.00064983877,0.0014136726,0.0059381695,0.001122924,0.0020370153,0.0017348925,0.0028495668,0.002602752],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010856193,0.00019688513,0.0050699343,0.00054770854,0.00021869777,0.00014193985,0.00027088315,0.064112574,0.0011148783,0.029373735,0.00926597,0.8895782],"study_design_scores_gemma":[0.000045055236,0.000058948335,0.002709235,0.00018548536,0.00004552573,0.00014402105,0.0001691878,0.8471239,0.0013010295,0.1352072,0.012954277,0.000056175344],"about_ca_topic_score_codex":0.004608159,"about_ca_topic_score_gemma":0.0043359213,"teacher_disagreement_score":0.010534617,"about_ca_system_score_codex":0.0018567408,"about_ca_system_score_gemma":0.002361108,"threshold_uncertainty_score":0.055713058},"labels":[],"label_agreement":null},{"id":"W2398309575","doi":"10.3233/978-1-61499-564-7-1101","title":"Evaluating a Hierarchical Clinical Event Linkage Model for Clinic-Specific Databases","year":2015,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network","funders":"","keywords":"Computer science; Meaning (existential); Linkage (software); Event (particle physics); Relational database; Database; Hierarchical database model; Information retrieval; Relational model; Data mining; Natural language processing; Psychology","score_opus":0.44042628285241064,"score_gpt":0.5519080042456707,"score_spread":0.11148172139326001,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398309575","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36209822,0.00046158853,0.62282145,0.0016156735,0.00006326498,0.0012630043,0.0037858416,0.003990706,0.0039002565],"genre_scores_gemma":[0.5614663,0.00025178207,0.4328343,0.00017062954,0.00002036197,0.00048771698,0.0037472637,0.00012654356,0.0008952268],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99250835,0.00371961,0.0009513873,0.0008485151,0.0016852898,0.00028685672],"domain_scores_gemma":[0.9668251,0.025765972,0.0015464302,0.002333879,0.0030838891,0.00044465993],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017998228,0.0005468378,0.0006879007,0.0023331903,0.0009811401,0.0036194457,0.0019445023,0.0011744737,0.0016789209],"category_scores_gemma":[0.04429015,0.00037144366,0.0014587201,0.0028084994,0.00076710887,0.004235274,0.001950253,0.00075415266,0.0003579717],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018856967,0.000931094,0.036559287,0.0005345534,0.00037161662,0.00032898426,0.0009310035,0.73162097,0.0027608895,0.056871295,0.0040225144,0.16318202],"study_design_scores_gemma":[0.0000654395,0.0002123775,0.0016641556,0.00004502556,0.000079341895,0.00008666432,0.00033589418,0.9857757,0.0013608543,0.008533923,0.0018097466,0.000030958276],"about_ca_topic_score_codex":0.039964482,"about_ca_topic_score_gemma":0.025967691,"teacher_disagreement_score":0.039964482,"about_ca_system_score_codex":0.0044861096,"about_ca_system_score_gemma":0.005205811,"threshold_uncertainty_score":0.09518492},"labels":[],"label_agreement":null},{"id":"W2398842345","doi":"","title":"Benchmarking infrastructure for mutation text mining.","year":2012,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Benchmarking; SPARQL; Information retrieval; Ontology; Data mining; RDF; Annotation; Benchmark (surveying); Data science; Semantic Web; Artificial intelligence","score_opus":0.013068377181528031,"score_gpt":0.27889584155909825,"score_spread":0.2658274643775702,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398842345","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046847735,0.001484668,0.5297846,0.002272817,0.00037003282,0.003631276,0.07245508,0.31852758,0.024626195],"genre_scores_gemma":[0.21374963,0.000753589,0.39510182,0.00068885187,0.00012271688,0.0035929407,0.36383882,0.016896706,0.005254917],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98313373,0.004984731,0.0035566757,0.0025723043,0.0048093516,0.0009432229],"domain_scores_gemma":[0.959194,0.010877548,0.0041308985,0.01219252,0.011600708,0.0020042427],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021247327,0.001958516,0.0013845745,0.00974378,0.001991519,0.003506613,0.0046696397,0.0016589771,0.00681335],"category_scores_gemma":[0.05260025,0.00083551806,0.0013357118,0.0103848735,0.0009831899,0.009068016,0.0052052527,0.0017316303,0.0050774636],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023347153,0.0017413578,0.025527002,0.00514969,0.0005388017,0.0013639809,0.0021294805,0.045711387,0.028607752,0.06143092,0.30708018,0.51838475],"study_design_scores_gemma":[0.0005030838,0.0009838478,0.026526866,0.0011150348,0.00031615558,0.0018358098,0.0014287636,0.42771184,0.08627653,0.08588517,0.36691803,0.0004989651],"about_ca_topic_score_codex":0.0066610486,"about_ca_topic_score_gemma":0.004614929,"teacher_disagreement_score":0.021247327,"about_ca_system_score_codex":0.003376224,"about_ca_system_score_gemma":0.0058077937,"threshold_uncertainty_score":0.11236793},"labels":[],"label_agreement":null},{"id":"W2398915717","doi":"","title":"Bio2RDF release 3: a larger connected network of linked data for the life sciences","year":2014,"lang":"en","type":"article","venue":"Research Publications (Maastricht University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":51,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"SPARQL; Computer science; Linked data; Data quality; Rest (music); Information retrieval; Database; World Wide Web; RDF; Data science; Data mining; Semantic Web; Engineering","score_opus":0.17671124927128493,"score_gpt":0.37344403281622385,"score_spread":0.19673278354493892,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2398915717","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056750197,0.000891391,0.04961584,0.00178622,0.0006955368,0.0008392632,0.876096,0.054926097,0.009474532],"genre_scores_gemma":[0.0071640303,0.00034116113,0.030537423,0.00030123215,0.000055233224,0.00070268154,0.9541863,0.004333371,0.00237853],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9953949,0.00081569015,0.00059705233,0.000953017,0.0019009602,0.00033835543],"domain_scores_gemma":[0.98638874,0.004194873,0.0009844991,0.004188193,0.0029901438,0.0012535166],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010281258,0.001979628,0.0015848844,0.0057182675,0.0016855941,0.0051291445,0.003102061,0.0020600148,0.029483302],"category_scores_gemma":[0.031446546,0.0011166082,0.0020495816,0.0060419096,0.00063319545,0.0049808845,0.0058512725,0.0024165576,0.018168658],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012867752,0.00015119415,0.0052304105,0.0017882013,0.0003510304,0.00057301053,0.0005204501,0.004300129,0.008021316,0.013218218,0.9121007,0.05245861],"study_design_scores_gemma":[0.00040134526,0.00010789645,0.0068808594,0.00052458706,0.00013864474,0.00037125737,0.00022186217,0.0076747513,0.0070313895,0.013553009,0.9629267,0.00016762775],"about_ca_topic_score_codex":0.022497663,"about_ca_topic_score_gemma":0.016210916,"teacher_disagreement_score":0.029483302,"about_ca_system_score_codex":0.001955814,"about_ca_system_score_gemma":0.0055511147,"threshold_uncertainty_score":0.09863144},"labels":[],"label_agreement":null},{"id":"W2399216429","doi":"","title":"BioMixer: Visualizing Mappings of Biomedical Ontologies.","year":2012,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ontology; Ontology components; Open Biomedical Ontologies; IDEF5; Computer science; Visualization; Process ontology; Information retrieval; Upper ontology; Data science; Ontology alignment; Semantic Web; Data mining; Epistemology","score_opus":0.023941269654700035,"score_gpt":0.3056664902052815,"score_spread":0.28172522055058147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2399216429","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016517607,0.00461229,0.80580974,0.0034880461,0.0005666549,0.00038767932,0.031807747,0.1135758,0.023234429],"genre_scores_gemma":[0.12513646,0.0055084773,0.81597704,0.0009577202,0.00015164097,0.00082958385,0.033293825,0.008523299,0.009621946],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9993705,0.00020737671,0.000073591254,0.000091744725,0.00021759873,0.000039157334],"domain_scores_gemma":[0.9974789,0.0015888646,0.00026174335,0.0002348659,0.00029431208,0.00014137203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021099558,0.0013249359,0.00062553625,0.004940528,0.0007129486,0.0025971641,0.0010423125,0.0013581093,0.017652167],"category_scores_gemma":[0.007916788,0.00049380196,0.0008760134,0.0037074308,0.000461836,0.0037558135,0.0033880458,0.0014961517,0.0025110312],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009390024,0.00022589299,0.0076174247,0.0051108375,0.00052677776,0.0017585151,0.0103622,0.0154183125,0.046737563,0.121319234,0.3415778,0.44840643],"study_design_scores_gemma":[0.00019231885,0.000105214036,0.006446263,0.0009821239,0.00020622,0.0022617711,0.0018749429,0.09767721,0.03545869,0.15188989,0.70269847,0.00020674834],"about_ca_topic_score_codex":0.0041357884,"about_ca_topic_score_gemma":0.005727262,"teacher_disagreement_score":0.017652167,"about_ca_system_score_codex":0.0006513188,"about_ca_system_score_gemma":0.0011237952,"threshold_uncertainty_score":0.059052348},"labels":[],"label_agreement":null},{"id":"W2400064635","doi":"","title":"Molecular Symmetry and Specialization of Atomic Connectivity by Class-based Reasoning of Chemical Structure.","year":2012,"lang":"en","type":"article","venue":"Research Publications (Maastricht University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Computer science; Representation (politics); Context (archaeology); Class (philosophy); RDF; Semantic Web; Theoretical computer science; Artificial intelligence; Information retrieval; Biology","score_opus":0.023430702370270874,"score_gpt":0.2981535510126697,"score_spread":0.2747228486423988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2400064635","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023720862,0.00034787727,0.966708,0.0009435746,0.000068759626,0.00018451217,0.0013635333,0.0011890624,0.005473747],"genre_scores_gemma":[0.2894709,0.00056535617,0.70315725,0.00032958746,0.00006588211,0.00017078548,0.0047421325,0.00020617446,0.0012919158],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99724776,0.0009102468,0.00030621394,0.00046587255,0.0009289649,0.00014096692],"domain_scores_gemma":[0.9929275,0.004391687,0.000659429,0.0013220888,0.0005248479,0.00017443109],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035336474,0.0005210295,0.00055381225,0.003968568,0.0009231559,0.0027119827,0.002060536,0.0010030995,0.0031542778],"category_scores_gemma":[0.014935063,0.00049300777,0.0031080581,0.0021455665,0.002085989,0.008732728,0.002045606,0.0016186121,0.0006110304],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002143235,0.00021138361,0.012293265,0.0005806218,0.00044152784,0.0010576692,0.0017766985,0.043449484,0.00947619,0.69903916,0.011345538,0.22011413],"study_design_scores_gemma":[0.000051948602,0.00004129383,0.0025196308,0.0001784604,0.00028365775,0.0006350606,0.000553009,0.32453218,0.010568977,0.63087,0.029711135,0.000054750635],"about_ca_topic_score_codex":0.010668558,"about_ca_topic_score_gemma":0.012653626,"teacher_disagreement_score":0.010668558,"about_ca_system_score_codex":0.0014144267,"about_ca_system_score_gemma":0.0019414797,"threshold_uncertainty_score":0.021212876},"labels":[],"label_agreement":null},{"id":"W2401831244","doi":"","title":"A primer of E-health terms.","year":2005,"lang":"es","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PricewaterhouseCoopers (Canada)","funders":"","keywords":"Primer (cosmetics); Computer science; Chemistry","score_opus":0.02327852606854828,"score_gpt":0.2695176484089844,"score_spread":0.24623912234043613,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401831244","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00222651,0.08720848,0.65388286,0.11465722,0.008190811,0.0011387251,0.0356929,0.010623179,0.08637935],"genre_scores_gemma":[0.015355618,0.043614414,0.86582124,0.024101486,0.0031150586,0.001439467,0.018013418,0.0016003713,0.0269389],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99691147,0.0013741947,0.00077837956,0.00027460814,0.00053250825,0.00012879698],"domain_scores_gemma":[0.977219,0.017761938,0.0010527018,0.0010429454,0.0020110356,0.000912303],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006748655,0.00095212506,0.0010956147,0.01564392,0.0015489521,0.0061086547,0.0022888638,0.0032855696,0.032091063],"category_scores_gemma":[0.0210134,0.0010533946,0.0014758861,0.010894494,0.0023623114,0.014310673,0.004305416,0.005484151,0.018563466],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016278647,0.00008198426,0.00068625336,0.004662212,0.000119145065,0.00093033013,0.0019066065,0.00066488073,0.0028924393,0.39617187,0.31074512,0.2809764],"study_design_scores_gemma":[0.000019982108,0.000014038549,0.00024298266,0.0012161402,0.000027188107,0.0007361822,0.00031377695,0.00054427946,0.00043395188,0.063241266,0.9331883,0.000021899054],"about_ca_topic_score_codex":0.0018749688,"about_ca_topic_score_gemma":0.004108922,"teacher_disagreement_score":0.032091063,"about_ca_system_score_codex":0.0013277214,"about_ca_system_score_gemma":0.002919716,"threshold_uncertainty_score":0.10735524},"labels":[],"label_agreement":null},{"id":"W2401857473","doi":"10.1101/055525","title":"Knowledge.Bio: A Web Application for Exploring, Building and Sharing Webs of Biomedical Relationships Mined from PubMed","year":2016,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Eagle Ridge Hospital","funders":"National Center for Advancing Translational Sciences; European Commission; Leids Universitair Medisch Centrum; Nederlandse Organisatie voor Wetenschappelijk Onderzoek; National Institutes of Health; Universiteit Leiden; Georgia Clinical and Translational Science Alliance","keywords":"World Wide Web; Computer science; Interface (matter); Information retrieval; Knowledge sharing; Data science; Knowledge management","score_opus":0.04667692690815123,"score_gpt":0.26088413938866567,"score_spread":0.21420721248051444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2401857473","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012922041,0.005119134,0.2516143,0.0034852047,0.0004535775,0.0014968194,0.27701834,0.41177678,0.036113743],"genre_scores_gemma":[0.06011132,0.0067919213,0.55952305,0.002107864,0.00036878328,0.0030230235,0.31770056,0.031584933,0.018788552],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989973,0.00021124035,0.00016095518,0.00021949522,0.00036070048,0.00005029226],"domain_scores_gemma":[0.9959181,0.002539353,0.00040886365,0.00055934064,0.00027934456,0.0002949465],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031952504,0.002316869,0.0013718794,0.013219681,0.0011946555,0.0043403413,0.0018265131,0.0016429668,0.049728353],"category_scores_gemma":[0.011532813,0.0011911297,0.0013754363,0.0075765746,0.00061045703,0.0048051146,0.005181925,0.0013867181,0.02283305],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011696622,0.00027812386,0.006222779,0.011575491,0.001216605,0.0023067594,0.0022836353,0.0036307345,0.015946012,0.02757728,0.60190606,0.3258869],"study_design_scores_gemma":[0.0006239539,0.00016718441,0.01056751,0.0021773416,0.00047594486,0.0027247025,0.0009401394,0.034859672,0.01208767,0.07599974,0.8590019,0.0003742092],"about_ca_topic_score_codex":0.0026223673,"about_ca_topic_score_gemma":0.0060521266,"teacher_disagreement_score":0.049728353,"about_ca_system_score_codex":0.00081053714,"about_ca_system_score_gemma":0.0021782583,"threshold_uncertainty_score":0.16635787},"labels":[],"label_agreement":null},{"id":"W2402464105","doi":"","title":"LinkedCT Live: Platform for Online Curation of Clinical Trials Data.","year":2015,"lang":"en","type":"article","venue":"International Semantic Web Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Data curation; Computer science; Clinical trial; World Wide Web; Data science; Linked data; Open data; Web application; Semantic Web; Bioinformatics","score_opus":0.536467777306551,"score_gpt":0.5289484315586915,"score_spread":0.007519345747859507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402464105","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019854316,0.006767475,0.22213687,0.008912984,0.0020062681,0.005835381,0.6097886,0.11066145,0.031905614],"genre_scores_gemma":[0.017189365,0.004834955,0.21882647,0.007936846,0.0010540339,0.009131003,0.713941,0.015993956,0.011092445],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.980923,0.008178735,0.0038384122,0.0017481912,0.004783329,0.00052835455],"domain_scores_gemma":[0.8846125,0.070094675,0.011669325,0.01637514,0.010319182,0.0069291666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.031334653,0.0019619248,0.0024790883,0.014942298,0.0016355662,0.009044526,0.004324626,0.00370655,0.060037248],"category_scores_gemma":[0.12705873,0.001571553,0.0025800585,0.013676691,0.0011487615,0.005684516,0.011235134,0.0056703785,0.028864743],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016006904,0.00028791043,0.0028220823,0.011054563,0.0013014929,0.0006840312,0.00084226625,0.001974685,0.005668717,0.026085628,0.8234936,0.124184385],"study_design_scores_gemma":[0.0012009401,0.00015460487,0.004499731,0.0029268118,0.000464588,0.00044966742,0.00017092633,0.0041498356,0.0057010227,0.04840583,0.93164074,0.00023538798],"about_ca_topic_score_codex":0.007254023,"about_ca_topic_score_gemma":0.008976058,"teacher_disagreement_score":0.060037248,"about_ca_system_score_codex":0.0027030355,"about_ca_system_score_gemma":0.019222284,"threshold_uncertainty_score":0.20084465},"labels":[],"label_agreement":null},{"id":"W2402799147","doi":"10.3233/978-1-61499-564-7-1110","title":"A Pilot Ontology for Healthcare Quality Indicators","year":2015,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Ontology; Quality (philosophy); Health care; Coding (social sciences); Computer science; Set (abstract data type); Data science; Process management; Knowledge management; Business; Political science","score_opus":0.15909382839470915,"score_gpt":0.4516513110756969,"score_spread":0.29255748268098775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402799147","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08203291,0.0004421748,0.7729592,0.004297798,0.00045755162,0.00432838,0.09571048,0.017453566,0.022318078],"genre_scores_gemma":[0.12620522,0.00043528568,0.79775566,0.00075698394,0.00005334162,0.0017197282,0.067118935,0.0013451608,0.004609703],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9979711,0.000385873,0.00053221436,0.00042134157,0.00055794156,0.0001314152],"domain_scores_gemma":[0.99333876,0.0030716371,0.0004854972,0.0009462909,0.0017581582,0.0003996235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002966072,0.00060261914,0.0005450521,0.0031895207,0.0012695737,0.0019134614,0.0010764668,0.0010700282,0.0056753005],"category_scores_gemma":[0.008643895,0.00044479826,0.002099625,0.003489297,0.000665788,0.004719068,0.0018434279,0.0013951987,0.0010198108],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012302912,0.0018168885,0.035622526,0.005752662,0.00037948746,0.0032960116,0.007695483,0.032221247,0.06295518,0.22948234,0.16586675,0.45368126],"study_design_scores_gemma":[0.0002994651,0.00032795093,0.021680377,0.00083587767,0.00033930928,0.002026227,0.002820138,0.08579486,0.021171449,0.048124302,0.8163135,0.00026664685],"about_ca_topic_score_codex":0.021488301,"about_ca_topic_score_gemma":0.02178515,"teacher_disagreement_score":0.021488301,"about_ca_system_score_codex":0.0027268713,"about_ca_system_score_gemma":0.0054621855,"threshold_uncertainty_score":0.042726457},"labels":[],"label_agreement":null},{"id":"W2402912048","doi":"10.3233/978-1-61499-432-9-1125","title":"Addressing the Challenge of Encoding Causal Epidemiological Knowledge in Formal Ontologies: A Practical Perspective","year":2014,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University Health Centre","funders":"","keywords":"Perspective (graphical); Computer science; Encoding (memory); Data science; Knowledge management; Population; Public health; Formal description; Artificial intelligence; Medicine; Environmental health; Programming language; Pathology","score_opus":0.2107657806240874,"score_gpt":0.47029901122893664,"score_spread":0.25953323060484923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2402912048","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019931586,0.000829266,0.97747076,0.016248142,0.00014552522,0.000091205984,0.00037351734,0.0003375356,0.0025108033],"genre_scores_gemma":[0.10387757,0.0031372146,0.8879261,0.002097969,0.0005009106,0.00024549782,0.0009167033,0.0001595247,0.0011385923],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9799604,0.012435056,0.0022344827,0.0012900395,0.003512694,0.0005673271],"domain_scores_gemma":[0.89579827,0.08460241,0.0048025963,0.009642813,0.004239986,0.00091394584],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03394138,0.0011981067,0.0017801193,0.005105633,0.002773213,0.014368486,0.0054624444,0.005116789,0.002990745],"category_scores_gemma":[0.08318567,0.0014946659,0.0026847275,0.007497857,0.009136842,0.026473166,0.008958442,0.008045345,0.00075902254],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000033851116,0.000054329736,0.00061499345,0.0006066359,0.000103787424,0.00040827726,0.0012955723,0.021509884,0.00048215466,0.9330217,0.0027595535,0.03910924],"study_design_scores_gemma":[0.000014238148,0.0000148988975,0.0000761995,0.00029455134,0.000048703678,0.00019064643,0.0008333735,0.027122932,0.00067417655,0.9481908,0.022514088,0.000025356838],"about_ca_topic_score_codex":0.00855074,"about_ca_topic_score_gemma":0.006094531,"teacher_disagreement_score":0.9660586,"about_ca_system_score_codex":0.004058223,"about_ca_system_score_gemma":0.006189337,"threshold_uncertainty_score":0.1795013},"labels":[],"label_agreement":null},{"id":"W2404114287","doi":"10.1080/03155986.2005.11732731","title":"Development of a Decision Algorithm to Support Emergency Triage of Scrotal Pain and its Implementation in the met system","year":2005,"lang":"en","type":"article","venue":"INFOR Information Systems and Operational Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Children's Hospital of Eastern Ontario; University of Ottawa","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Triage; Workflow; Computer science; Scrotal Pain; Algorithm; Clinical decision support system; Decision support system; USable; Medical emergency; Medicine; Artificial intelligence; Scrotum","score_opus":0.05964374840539965,"score_gpt":0.40982565307998153,"score_spread":0.3501819046745819,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2404114287","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057372473,0.00018908024,0.91405874,0.00063537766,0.00010861714,0.00073233136,0.0013519085,0.023215154,0.0023362748],"genre_scores_gemma":[0.19507834,0.00010376939,0.80121964,0.0001724571,0.000044102577,0.00034768687,0.0016813753,0.00019353535,0.0011590407],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989538,0.00024512009,0.00022732142,0.0002465749,0.00025760685,0.00006953963],"domain_scores_gemma":[0.9973889,0.0014252469,0.00015832816,0.00020334158,0.0006843048,0.00013995536],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017803285,0.0004674227,0.0008012654,0.001443692,0.000524982,0.0020686137,0.0011524397,0.0011598797,0.0030877215],"category_scores_gemma":[0.007015887,0.0004403071,0.0006617605,0.0008690432,0.0002575581,0.0017219842,0.0006759475,0.0007580593,0.001202295],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019057272,0.00086081045,0.019438341,0.0007709408,0.00035659855,0.0015434829,0.0008382061,0.1273413,0.045710273,0.014736335,0.021870585,0.7646274],"study_design_scores_gemma":[0.00023147916,0.00020833791,0.0023290955,0.000077641766,0.00008075168,0.00046563384,0.00014427814,0.95708084,0.02025369,0.003574864,0.015495321,0.000058053083],"about_ca_topic_score_codex":0.0023357682,"about_ca_topic_score_gemma":0.0018092954,"teacher_disagreement_score":0.0030877215,"about_ca_system_score_codex":0.00060419005,"about_ca_system_score_gemma":0.0012427268,"threshold_uncertainty_score":0.010329485},"labels":[],"label_agreement":null},{"id":"W2404340837","doi":"","title":"Semantic search for the life sciences.","year":2011,"lang":"fr","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Semantic search; Computer science; Information retrieval; Data science; World Wide Web; Semantic Web","score_opus":0.05373632683197375,"score_gpt":0.28198533466662085,"score_spread":0.2282490078346471,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2404340837","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023655219,0.07733036,0.7184407,0.01742814,0.0024636574,0.00072066224,0.07093905,0.018888546,0.070133634],"genre_scores_gemma":[0.19076112,0.028823454,0.63807994,0.0025794825,0.0007355101,0.000743017,0.113847405,0.000937478,0.02349257],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998803,0.0003531752,0.00018778272,0.00018716198,0.00040207937,0.00006680853],"domain_scores_gemma":[0.99835193,0.0007946243,0.00013133015,0.00031184318,0.000291127,0.00011923083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018533752,0.00077699975,0.001223162,0.0065429807,0.0012502832,0.0032716696,0.0012187178,0.0016739594,0.0111411875],"category_scores_gemma":[0.006958703,0.000377147,0.0015348457,0.007970356,0.001047213,0.0071982364,0.0025912542,0.0011616749,0.0069046724],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040530454,0.00015817405,0.0016385962,0.0038262284,0.00032363582,0.0005127978,0.00050643145,0.005586347,0.0043904088,0.46114326,0.14850806,0.37300077],"study_design_scores_gemma":[0.00006351662,0.000049324528,0.0014580819,0.00070155866,0.00016997491,0.00047339784,0.00049076945,0.03539159,0.0029970699,0.6306022,0.3275655,0.000037120146],"about_ca_topic_score_codex":0.005508174,"about_ca_topic_score_gemma":0.010776691,"teacher_disagreement_score":0.0111411875,"about_ca_system_score_codex":0.0017001007,"about_ca_system_score_gemma":0.0030822176,"threshold_uncertainty_score":0.037270963},"labels":[],"label_agreement":null},{"id":"W2405371753","doi":"10.3233/978-1-61499-438-1-409","title":"The Cardiovascular Disease Ontology","year":2014,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Hospitalier Universitaire de Sherbrooke","funders":"","keywords":"Disease; Ontology; Computer science; Medicine; Internal medicine; Philosophy; Epistemology","score_opus":0.02814763827631525,"score_gpt":0.2623300727598922,"score_spread":0.23418243448357692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405371753","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008463662,0.03861879,0.33209157,0.025664363,0.0057098726,0.00043328843,0.019332377,0.0023689435,0.56731707],"genre_scores_gemma":[0.1390923,0.06798508,0.5136899,0.018632593,0.0038941936,0.0008452356,0.043597564,0.0011049549,0.21115817],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99953246,0.00008535058,0.000058607802,0.00008941312,0.00018680026,0.000047395195],"domain_scores_gemma":[0.9995177,0.00019123493,0.000037850456,0.000054520042,0.00011219962,0.00008647191],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00068890996,0.0003510624,0.00033718676,0.0031396912,0.0012222833,0.0025290933,0.00085430325,0.0011519761,0.008673591],"category_scores_gemma":[0.001736287,0.00025569054,0.00055850577,0.0029983532,0.0011341852,0.0034557036,0.0020642616,0.0017068345,0.0026159673],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018728864,0.000018042936,0.00057795155,0.0004162174,0.000012403626,0.0002784193,0.00062566303,0.0006273761,0.0007868729,0.7192447,0.1092382,0.16815546],"study_design_scores_gemma":[0.0000034831942,0.000002580834,0.0004295473,0.00018658585,0.0000062264185,0.0004213082,0.00010813663,0.00041844597,0.00008384497,0.05543404,0.94289863,0.00000716006],"about_ca_topic_score_codex":0.007767101,"about_ca_topic_score_gemma":0.007908014,"teacher_disagreement_score":0.008673591,"about_ca_system_score_codex":0.0017310121,"about_ca_system_score_gemma":0.003517858,"threshold_uncertainty_score":0.029016018},"labels":[],"label_agreement":null},{"id":"W2405515544","doi":"","title":"Search filter precision can be improved by NOTing out irrelevant content.","year":2011,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"PsycINFO; CINAHL; Computer science; Information retrieval; MEDLINE; Filter (signal processing); Search engine indexing; Computer vision","score_opus":0.10625459839194751,"score_gpt":0.26324094073440724,"score_spread":0.15698634234245973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2405515544","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3059777,0.11832515,0.47993317,0.021623844,0.002536496,0.013104809,0.014680337,0.005096914,0.038721606],"genre_scores_gemma":[0.6137642,0.012256215,0.3511386,0.006151559,0.0011739207,0.008191662,0.0051859445,0.00042166555,0.0017162078],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.64306515,0.20080882,0.089363866,0.013791716,0.051254813,0.0017156489],"domain_scores_gemma":[0.07469798,0.85989374,0.027291799,0.021877678,0.015752664,0.0004861022],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.4074348,0.0020841078,0.0046823514,0.032868832,0.0021974037,0.009357393,0.003553897,0.0032034973,0.0028243032],"category_scores_gemma":[0.74327314,0.001519062,0.0057729455,0.023003977,0.003074035,0.010154962,0.004374498,0.0017031351,0.0011828543],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0037055623,0.0004891011,0.16901115,0.05888134,0.010860915,0.0005122315,0.009326366,0.004195398,0.0055408524,0.012728813,0.0204476,0.7043007],"study_design_scores_gemma":[0.003774055,0.004263554,0.39609024,0.058776837,0.052992553,0.004049831,0.010303096,0.047591284,0.04704439,0.18012361,0.19331929,0.0016712844],"about_ca_topic_score_codex":0.0028451695,"about_ca_topic_score_gemma":0.0043662232,"teacher_disagreement_score":0.5925652,"about_ca_system_score_codex":0.00420209,"about_ca_system_score_gemma":0.007342173,"threshold_uncertainty_score":0.7307384},"labels":[],"label_agreement":null},{"id":"W2406086978","doi":"10.3233/978-1-61499-289-9-1207","title":"PHIO: A Knowledge Base for Interpretation and Calculation of Public Health Indicators","year":2013,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University; McGill University Health Centre","funders":"","keywords":"Ontology; Public health; Interpretation (philosophy); Data science; Population; Knowledge base; Knowledge management; Population health; Health indicator; Computer science; Environmental health; Medicine; World Wide Web; Nursing","score_opus":0.05353176952629964,"score_gpt":0.37783510313264684,"score_spread":0.3243033336063472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2406086978","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01073828,0.00081102055,0.85278046,0.0021018954,0.00025269165,0.0017941633,0.08086587,0.02216562,0.028490014],"genre_scores_gemma":[0.06580908,0.001835405,0.83388615,0.0006376408,0.00007965955,0.001133881,0.091849454,0.0011302351,0.0036384596],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99752444,0.0004198108,0.00056038843,0.00045298788,0.00090734847,0.00013495525],"domain_scores_gemma":[0.99265,0.0031421485,0.00064998347,0.0015518254,0.0017179599,0.0002880546],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004487899,0.0010851434,0.0010027665,0.01180157,0.0014094481,0.005551473,0.0025374407,0.0011369972,0.005650344],"category_scores_gemma":[0.018917188,0.0005969438,0.0015732754,0.009017778,0.0009686249,0.005891387,0.003939413,0.0015583505,0.0022945504],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028388295,0.0004994942,0.015911898,0.004557201,0.0005756835,0.002491859,0.003891124,0.031283177,0.008879113,0.19951671,0.13116015,0.6009497],"study_design_scores_gemma":[0.000091296795,0.00006875912,0.00961158,0.002325186,0.00039755309,0.00073162746,0.0023090395,0.11375497,0.011119347,0.12838721,0.7309919,0.00021146178],"about_ca_topic_score_codex":0.044294786,"about_ca_topic_score_gemma":0.044687517,"teacher_disagreement_score":0.044294786,"about_ca_system_score_codex":0.0030752777,"about_ca_system_score_gemma":0.008594188,"threshold_uncertainty_score":0.08807391},"labels":[],"label_agreement":null},{"id":"W2406405072","doi":"10.3233/978-1-61499-432-9-1075","title":"Frame Semantics-based Study of Verbs across Medical Genres","year":2014,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Université du Québec à Montréal","funders":"","keywords":"FrameNet; Computer science; Natural language processing; Semantics (computer science); Medical terminology; Terminology; Frame (networking); Artificial intelligence; Linguistics; Field (mathematics); Programming language; Parsing; Mathematics","score_opus":0.03303768602480658,"score_gpt":0.39872855321693806,"score_spread":0.3656908671921315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2406405072","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90555626,0.00071380235,0.07625096,0.000439855,0.00009456403,0.00026813857,0.001674989,0.00022289438,0.014778428],"genre_scores_gemma":[0.97153646,0.000113733135,0.024863489,0.000051528466,0.00003645402,0.00020347627,0.0020549083,0.00010405632,0.001035995],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99740607,0.001374231,0.00017592547,0.00057841535,0.0003177343,0.00014764254],"domain_scores_gemma":[0.97948897,0.01615729,0.001436285,0.00078568136,0.0016842891,0.00044743513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031067007,0.00040839042,0.00042381344,0.0057272264,0.0014530083,0.0021303135,0.000613002,0.00096787454,0.0030968264],"category_scores_gemma":[0.017724331,0.00022826518,0.00070487306,0.0035095701,0.0022460693,0.0031454142,0.0014287184,0.0009138176,0.00043452386],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034706334,0.0010587536,0.1625703,0.0028877188,0.00038326334,0.003525589,0.1314215,0.00827223,0.19038716,0.23636682,0.007921492,0.25173458],"study_design_scores_gemma":[0.00042028195,0.0011422061,0.48237365,0.000917483,0.0004896753,0.006243545,0.07747702,0.15919273,0.044388693,0.12732379,0.09972646,0.0003044598],"about_ca_topic_score_codex":0.0045485594,"about_ca_topic_score_gemma":0.004106851,"teacher_disagreement_score":0.0057272264,"about_ca_system_score_codex":0.0018024791,"about_ca_system_score_gemma":0.00059667195,"threshold_uncertainty_score":0.01642996},"labels":[],"label_agreement":null},{"id":"W2407238648","doi":"","title":"Clinical guideline-driven personalized self-management diary for paediatric cancer survivors.","year":2014,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Izaak Walton Killam Health Centre; Dalhousie University","funders":"","keywords":"Guideline; Ontology; Knowledge base; Process (computing); Clinical Practice; Medicine; Computer science; Knowledge management; Nursing; Artificial intelligence","score_opus":0.024789206492859147,"score_gpt":0.30431544738480587,"score_spread":0.2795262408919467,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2407238648","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22849971,0.006251542,0.5509364,0.021653991,0.0015178212,0.0076073306,0.10857523,0.030675706,0.044282343],"genre_scores_gemma":[0.29996988,0.0024808752,0.65226245,0.0012867493,0.00014232052,0.0018370502,0.03615817,0.00064232235,0.0052202484],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99797314,0.0008631092,0.00053419755,0.00022458269,0.00035401565,0.000050987124],"domain_scores_gemma":[0.989689,0.005742743,0.0013884224,0.00128508,0.0014276664,0.00046712754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046942574,0.0003489,0.00040084275,0.0017872723,0.00039834314,0.0016417704,0.00087389967,0.00062970247,0.0030161347],"category_scores_gemma":[0.021897245,0.00025755173,0.0005478865,0.0013495269,0.0001546949,0.0018320839,0.0010142202,0.0008036371,0.0016816324],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060987385,0.00042454537,0.05048272,0.0020885302,0.00023544994,0.0011050837,0.0039147576,0.0063451803,0.00628542,0.0070091193,0.12597251,0.7955268],"study_design_scores_gemma":[0.00045262408,0.00072526396,0.059722178,0.0030954941,0.000642585,0.0036800131,0.0054561673,0.10298274,0.030500745,0.032031212,0.7603533,0.0003576391],"about_ca_topic_score_codex":0.0019642368,"about_ca_topic_score_gemma":0.0060434733,"teacher_disagreement_score":0.0046942574,"about_ca_system_score_codex":0.00079716544,"about_ca_system_score_gemma":0.0020340977,"threshold_uncertainty_score":0.024825871},"labels":[],"label_agreement":null},{"id":"W2407659378","doi":"","title":"Ecotoxicology Data Federation with SADI Semantic Web Services.","year":2012,"lang":"en","type":"article","venue":"SWAT4LS","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"SPARQL; World Wide Web; Computer science; Premise; Semantic Web; RDF; Data science; Information retrieval","score_opus":0.025816486318144635,"score_gpt":0.2859690024521057,"score_spread":0.26015251613396106,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2407659378","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07486529,0.00072327384,0.73786324,0.00454671,0.00043910107,0.001770831,0.020019084,0.112548895,0.047223464],"genre_scores_gemma":[0.4339515,0.0005875607,0.45692846,0.0016193347,0.000124739,0.0010706257,0.092126675,0.0035444144,0.010046616],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9949234,0.0010290907,0.00093568605,0.000760654,0.0019482502,0.00040297967],"domain_scores_gemma":[0.9888881,0.002398533,0.00081135985,0.005113875,0.0020886867,0.0006993849],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011591667,0.0006764843,0.0007963576,0.0022796765,0.0015162164,0.0037731796,0.0026838675,0.0011406188,0.0029711511],"category_scores_gemma":[0.01103627,0.0005877578,0.001973916,0.0026200265,0.0009698198,0.0060591386,0.005337086,0.0015217795,0.0019521865],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0047553442,0.0023801643,0.07118142,0.0028080088,0.0014206431,0.0027103482,0.0052990387,0.029360311,0.076151066,0.20773931,0.16582195,0.43037236],"study_design_scores_gemma":[0.000296903,0.00036396473,0.013711313,0.00039848938,0.00041363022,0.0016678083,0.0019447185,0.2931323,0.08080339,0.095432006,0.51158494,0.00025051172],"about_ca_topic_score_codex":0.007979131,"about_ca_topic_score_gemma":0.00720154,"teacher_disagreement_score":0.011591667,"about_ca_system_score_codex":0.0025046126,"about_ca_system_score_gemma":0.004736792,"threshold_uncertainty_score":0.061303377},"labels":[],"label_agreement":null},{"id":"W2408039204","doi":"","title":"A Physician's Authoring Tool for Generation of Personalized Health Education in Reconstructive Surgery.","year":2006,"lang":"en","type":"article","venue":"National Conference on Artificial Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Personalization; Domain (mathematical analysis); Computer science; Health care; Multimedia; World Wide Web","score_opus":0.1419393657558969,"score_gpt":0.3870617355454141,"score_spread":0.2451223697895172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408039204","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007441949,0.0006163829,0.8584253,0.0009419647,0.00053982047,0.0005053667,0.008170117,0.11090154,0.012457554],"genre_scores_gemma":[0.07234896,0.00060111785,0.8830854,0.0005039157,0.00012684765,0.0006642537,0.013487588,0.005776736,0.023405233],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991967,0.00025552243,0.000096316384,0.00012835198,0.0002902695,0.00003273895],"domain_scores_gemma":[0.9955468,0.0033422522,0.00016043644,0.00041469757,0.0003619403,0.0001738496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020036679,0.0010766581,0.000506741,0.0016557148,0.00050068606,0.0014723977,0.0010230162,0.0014795688,0.031857524],"category_scores_gemma":[0.0077911047,0.0005864562,0.00066860934,0.0009021346,0.00036360134,0.0012761944,0.0013830444,0.0008775373,0.007288439],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065042655,0.00034007672,0.003966261,0.0012178911,0.00009967576,0.0027255272,0.001964354,0.01116403,0.019068873,0.020723477,0.3164763,0.6216031],"study_design_scores_gemma":[0.00046600797,0.00022259438,0.0025926419,0.0005221962,0.000115729825,0.00349593,0.0004245949,0.21329431,0.035539914,0.033847094,0.70935905,0.000119842705],"about_ca_topic_score_codex":0.0007719201,"about_ca_topic_score_gemma":0.0019126456,"teacher_disagreement_score":0.031857524,"about_ca_system_score_codex":0.00039423458,"about_ca_system_score_gemma":0.0006786831,"threshold_uncertainty_score":0.106574},"labels":[],"label_agreement":null},{"id":"W2408224326","doi":"","title":"Biological event extraction using subgraph matching.","year":2010,"lang":"en","type":"article","venue":"Semantic Mining in Biomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Subgraph isomorphism problem; Dependency (UML); Event (particle physics); Task (project management); Matching (statistics); Dependency graph; Graph; Information extraction; Relationship extraction; Biomedical text mining; Artificial intelligence; Data mining; Natural language processing; Theoretical computer science; Text mining; Mathematics","score_opus":0.025152167136819528,"score_gpt":0.3264374321826622,"score_spread":0.3012852650458427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2408224326","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03395268,0.0023416972,0.9063677,0.00092442613,0.0002183826,0.0014169782,0.02898363,0.018255783,0.0075386986],"genre_scores_gemma":[0.1623051,0.0015874108,0.7449926,0.00034675712,0.00013964507,0.0007505169,0.08567498,0.0007937104,0.0034092437],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99825054,0.00036206734,0.00024696757,0.00054043956,0.0004964208,0.0001035408],"domain_scores_gemma":[0.9968207,0.0015924475,0.00046181743,0.00055024435,0.00047294365,0.00010188335],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013230244,0.0011334103,0.00092878164,0.009508621,0.00078511285,0.0011189092,0.0013402336,0.0012623528,0.0042780032],"category_scores_gemma":[0.0072268154,0.00044735594,0.0024081029,0.006756513,0.00047009674,0.0023226817,0.0017499061,0.00080784416,0.002575233],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006119609,0.00047835146,0.014940546,0.0041682883,0.0009969083,0.0025468417,0.00080772134,0.018721547,0.07361955,0.04130701,0.05834934,0.783452],"study_design_scores_gemma":[0.00023196217,0.00041341945,0.023444194,0.00051871384,0.0011691917,0.006247746,0.00086655084,0.37787426,0.10435307,0.2794462,0.20526615,0.00016853621],"about_ca_topic_score_codex":0.004132671,"about_ca_topic_score_gemma":0.0060051526,"teacher_disagreement_score":0.009508621,"about_ca_system_score_codex":0.00078077475,"about_ca_system_score_gemma":0.0016911486,"threshold_uncertainty_score":0.014311373},"labels":[],"label_agreement":null},{"id":"W2413149168","doi":"","title":"More scans, more scanners.","year":2005,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Health Information","funders":"","keywords":"Best practice; Medical physics; Business; Computer science; Medicine; Management; Economics","score_opus":0.015406977601207806,"score_gpt":0.2455852178345929,"score_spread":0.2301782402333851,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2413149168","genre_codex":"dataset","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005928955,0.010880693,0.011892343,0.0058534658,0.002119625,0.0003271143,0.8458795,0.052907992,0.06421028],"genre_scores_gemma":[0.041526657,0.010163611,0.08850973,0.008383896,0.0017719609,0.0011905604,0.752816,0.021195717,0.074441746],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99867356,0.00027164235,0.00028833677,0.00028541137,0.00036385038,0.00011723451],"domain_scores_gemma":[0.9877033,0.0065986435,0.0009694498,0.0017700085,0.0020810834,0.0008774427],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0013588448,0.0012836496,0.0013515012,0.014541414,0.0007849872,0.003827313,0.0014074683,0.0016945009,0.42424524],"category_scores_gemma":[0.017442059,0.0013504524,0.0012819335,0.014112699,0.0004731987,0.008490477,0.0021583645,0.0014757898,0.17096855],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029401417,0.00008826365,0.0018028792,0.0020438603,0.00011788386,0.0002731227,0.0001049329,0.00020341598,0.0028513346,0.0020261137,0.8618337,0.12836064],"study_design_scores_gemma":[0.00019609038,0.000032857046,0.0033257469,0.00074492453,0.00011137169,0.0013293673,0.00015682576,0.0007336441,0.0027552685,0.008215027,0.98232347,0.000075554744],"about_ca_topic_score_codex":0.0040567955,"about_ca_topic_score_gemma":0.011945244,"teacher_disagreement_score":0.42424524,"about_ca_system_score_codex":0.0006478016,"about_ca_system_score_gemma":0.0015759779,"threshold_uncertainty_score":0.8212443},"labels":[],"label_agreement":null},{"id":"W2414137297","doi":"","title":"Introduction to a searchable database of William Osler photographs.","year":2009,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"World Wide Web; Computer science; Database","score_opus":0.016797393542775072,"score_gpt":0.24942783938022955,"score_spread":0.23263044583745449,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2414137297","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031421199,0.07850915,0.398261,0.041327152,0.013480108,0.0029244742,0.18552631,0.06875191,0.20807786],"genre_scores_gemma":[0.010194299,0.05406781,0.60558796,0.01633612,0.0052840216,0.0025325103,0.09907741,0.011779619,0.19514018],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99855906,0.0003004238,0.00035530006,0.00023198377,0.00047619364,0.000077020755],"domain_scores_gemma":[0.991698,0.0043415106,0.0004454328,0.0007578219,0.0019800146,0.00077714754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002832119,0.0011531959,0.0010888812,0.01633288,0.0011374934,0.0053263465,0.0021525945,0.0012860175,0.13767196],"category_scores_gemma":[0.019454563,0.00093439437,0.0011116836,0.012515709,0.00079260976,0.009388239,0.0038119706,0.0019678138,0.098180614],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000064171094,0.000042924337,0.00022318443,0.0010138141,0.000021549879,0.00025028508,0.00012044571,0.00013450664,0.0012698796,0.011308183,0.79915565,0.18639536],"study_design_scores_gemma":[0.000011847105,0.000013616352,0.00032814752,0.0004496676,0.000009857028,0.00063662644,0.00007666927,0.00029116488,0.00046172575,0.0063372124,0.99136084,0.00002251578],"about_ca_topic_score_codex":0.003170093,"about_ca_topic_score_gemma":0.0068708477,"teacher_disagreement_score":0.13767196,"about_ca_system_score_codex":0.0009432633,"about_ca_system_score_gemma":0.0017519605,"threshold_uncertainty_score":0.4605586},"labels":[],"label_agreement":null},{"id":"W2417314476","doi":"10.1007/978-3-319-30148-8_7","title":"How to Develop a KOS to Serve Interdisciplinarity","year":2016,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Ambiguity; Phenomenon; sort; Epistemology; Management science; Computer science; Data science; Sociology; Engineering ethics; Cognitive science; Psychology; Engineering; Philosophy; Information retrieval","score_opus":0.025599139897492674,"score_gpt":0.2764268791262178,"score_spread":0.2508277392287251,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2417314476","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024617324,0.00046284124,0.9612972,0.004538363,0.0005027241,0.00021153345,0.00070283725,0.0034879846,0.026334919],"genre_scores_gemma":[0.012553653,0.0006554659,0.97255707,0.0007206097,0.000069955444,0.0001965063,0.0015719597,0.0010372342,0.010637567],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9976561,0.0007105235,0.00043877124,0.00032435992,0.0007155746,0.00015463494],"domain_scores_gemma":[0.9956715,0.0009954004,0.00016487356,0.00133363,0.0015115474,0.00032308768],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0057937633,0.0007624472,0.0006615606,0.0026323362,0.001231356,0.0059207967,0.0016051433,0.0014027148,0.008209103],"category_scores_gemma":[0.010405712,0.0007794336,0.0016050907,0.0025658135,0.0019938585,0.019496126,0.005323482,0.003539818,0.00859102],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003547812,0.00010938938,0.0010826843,0.0005215136,0.000047997913,0.0003064972,0.001578474,0.0019983936,0.005068984,0.6750323,0.050693713,0.26352453],"study_design_scores_gemma":[0.000013142048,0.000022034643,0.0004053276,0.00050165254,0.00005050475,0.0006313762,0.001012122,0.013333951,0.005085732,0.3108018,0.66809875,0.00004364399],"about_ca_topic_score_codex":0.003663197,"about_ca_topic_score_gemma":0.0038364758,"teacher_disagreement_score":0.99407923,"about_ca_system_score_codex":0.0012571621,"about_ca_system_score_gemma":0.003844492,"threshold_uncertainty_score":0.030640721},"labels":[],"label_agreement":null},{"id":"W2418383179","doi":"10.1093/database/baw090","title":"MET network in PubMed: a text-mined network visualization and curation system","year":2016,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"National Human Genome Research Institute; National Institutes of Health; National Institute of General Medical Sciences; Ministry of Science and Technology, Taiwan; European Molecular Biology Laboratory","keywords":"Computer science; Task (project management); Metastasis; Data curation; Visualization; World Wide Web; Process (computing); Cancer; Information retrieval; Data science; Bioinformatics; Medicine; Artificial intelligence; Biology","score_opus":0.014714083094430574,"score_gpt":0.25834995446435843,"score_spread":0.24363587136992787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2418383179","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01212881,0.005228307,0.16011953,0.0035687604,0.0005148229,0.0020308534,0.60175794,0.20248851,0.012162426],"genre_scores_gemma":[0.04015243,0.0044976287,0.528719,0.0009896293,0.00023527275,0.003404032,0.4102604,0.0068937363,0.0048478474],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99768925,0.00059635076,0.00068750593,0.00050343375,0.00044859308,0.00007481801],"domain_scores_gemma":[0.9889836,0.00557035,0.0017451401,0.0013983999,0.0018086754,0.0004940036],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0051020375,0.002015415,0.0014780894,0.018527497,0.0013503658,0.0032120934,0.0020354697,0.0014642077,0.019403819],"category_scores_gemma":[0.02020289,0.0009985422,0.0018445336,0.011112493,0.000374226,0.00478007,0.004310636,0.0010561303,0.00903887],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015957769,0.00021306935,0.011715519,0.02640141,0.0012078785,0.0026137752,0.0028903752,0.0049367803,0.02537403,0.020292187,0.6068251,0.295934],"study_design_scores_gemma":[0.0005258167,0.00026682354,0.010474706,0.002571042,0.0007997217,0.0016142441,0.00083631504,0.037264787,0.017190307,0.021327198,0.90686536,0.00026362843],"about_ca_topic_score_codex":0.004194766,"about_ca_topic_score_gemma":0.009580348,"teacher_disagreement_score":0.99489796,"about_ca_system_score_codex":0.0012491036,"about_ca_system_score_gemma":0.0039942428,"threshold_uncertainty_score":0.0649122},"labels":[],"label_agreement":null},{"id":"W2418468785","doi":"","title":"A method for verifying a vector-based text classification system.","year":2008,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Computer science; Search engine indexing; Java; Suite; Set (abstract data type); Lisp; Similarity (geometry); Data mining; Index (typography); Information retrieval; Vector space model; Artificial intelligence; Programming language","score_opus":0.07073376003347205,"score_gpt":0.29004516005036646,"score_spread":0.21931140001689442,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2418468785","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045494055,0.00016710165,0.9751873,0.000179447,0.00023142113,0.00044072108,0.0018475981,0.01542295,0.0019740253],"genre_scores_gemma":[0.048538078,0.00012338528,0.9429565,0.00010749485,0.00006183911,0.0007347587,0.0043684933,0.0007013417,0.0024081438],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9846297,0.0024292115,0.0033427475,0.0020074625,0.0072463923,0.00034448094],"domain_scores_gemma":[0.9755275,0.008453148,0.0018951704,0.0043643806,0.009372049,0.00038765089],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008642539,0.0011476728,0.00079610985,0.006084044,0.0014850985,0.003543256,0.0018640866,0.0012850441,0.010389952],"category_scores_gemma":[0.04844197,0.0005941307,0.0010850152,0.0041501042,0.0011818982,0.004596066,0.0026454634,0.0011911211,0.00795547],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048194872,0.00022714163,0.007952257,0.0010434474,0.0002001118,0.00030121754,0.0005925911,0.005698902,0.027855314,0.045056496,0.030198757,0.8803918],"study_design_scores_gemma":[0.0005025252,0.0009521453,0.012624229,0.000764819,0.00033913285,0.0036199014,0.0009293362,0.4590199,0.15819925,0.13275282,0.22992937,0.00036654162],"about_ca_topic_score_codex":0.0030616433,"about_ca_topic_score_gemma":0.0023953847,"teacher_disagreement_score":0.010389952,"about_ca_system_score_codex":0.0011387895,"about_ca_system_score_gemma":0.0033594624,"threshold_uncertainty_score":0.04570669},"labels":[],"label_agreement":null},{"id":"W2419248742","doi":"","title":"Coding of drugs used by respondents of the Canadian Study of Health and Aging.","year":2002,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Institute of Population and Public Health; Health Canada","funders":"","keywords":"Formulary; Coding (social sciences); Medicine; Diagnosis code; Family medicine; Environmental health; Statistics","score_opus":0.03525682118942535,"score_gpt":0.25417963651147624,"score_spread":0.21892281532205088,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2419248742","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5263051,0.008942979,0.0042578517,0.0060383677,0.00034710698,0.0025410915,0.41084668,0.00025423605,0.040466566],"genre_scores_gemma":[0.85777813,0.0074344636,0.012588654,0.001833837,0.000048722257,0.0016213887,0.11172932,0.00006124033,0.006904288],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99554104,0.0007573506,0.0009444176,0.00030055048,0.0020572029,0.0003994262],"domain_scores_gemma":[0.98379296,0.003139177,0.0035175327,0.00063656666,0.0080173435,0.0008965128],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039608045,0.00029307595,0.00034635677,0.007327775,0.0017030361,0.0008074601,0.0010265436,0.0003640727,0.0029441114],"category_scores_gemma":[0.02545716,0.00016203761,0.0005033586,0.014871115,0.0007058437,0.00042540193,0.00078481704,0.00039703233,0.00043424242],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003037554,0.00006493852,0.811183,0.0029799843,0.00015066208,0.0004761738,0.015062222,0.00032067669,0.0015489634,0.0026725442,0.06961807,0.09561906],"study_design_scores_gemma":[0.000017180088,0.000029103803,0.9375053,0.00045217664,0.00008056619,0.00020499484,0.0072036036,0.00021046204,0.0005411311,0.0002387883,0.053473093,0.000043698106],"about_ca_topic_score_codex":0.9686155,"about_ca_topic_score_gemma":0.96859854,"teacher_disagreement_score":0.9686155,"about_ca_system_score_codex":0.014975039,"about_ca_system_score_gemma":0.03803985,"threshold_uncertainty_score":0.108652055},"labels":[],"label_agreement":null},{"id":"W2439277838","doi":"10.1080/17453054.2016.1182473","title":"A reusable anatomically segmented digital mannequin for public health communication","year":2016,"lang":"en","type":"article","venue":"Journal of Visual Communication in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Toronto","keywords":"Set (abstract data type); Computer science; Object (grammar); Data science; Multimedia; Data set; Human–computer interaction; World Wide Web; Visualization; Computer graphics (images); Artificial intelligence","score_opus":0.060473391531214137,"score_gpt":0.38850418322152414,"score_spread":0.32803079169031,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2439277838","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035993304,0.0008905045,0.7627001,0.0017841564,0.00032011344,0.0015954678,0.065904595,0.07426165,0.056550115],"genre_scores_gemma":[0.16176759,0.0012752396,0.6954441,0.0004882653,0.00013356554,0.0017487896,0.09541581,0.0069998526,0.036726728],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99960905,0.00007067249,0.000040831594,0.00007435181,0.0001698434,0.000035319077],"domain_scores_gemma":[0.9987571,0.00046967863,0.00009689187,0.00035700406,0.0002217082,0.00009764231],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007753633,0.0008709739,0.0004638262,0.0037415456,0.0008174825,0.0014632653,0.0010832013,0.0009615233,0.031243855],"category_scores_gemma":[0.0031155373,0.00031573387,0.0007810542,0.0025757528,0.0005037956,0.0018102325,0.002829357,0.000651556,0.011095414],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068873644,0.00021239961,0.0043278,0.0012696552,0.0000841958,0.002280062,0.0022925886,0.0067722495,0.051473413,0.023558876,0.1516455,0.7553947],"study_design_scores_gemma":[0.000071610775,0.0002089801,0.012109019,0.00058889214,0.00009378167,0.0045203567,0.0014708292,0.046433866,0.05813988,0.028153768,0.84801656,0.00019243926],"about_ca_topic_score_codex":0.0028159022,"about_ca_topic_score_gemma":0.002903296,"teacher_disagreement_score":0.031243855,"about_ca_system_score_codex":0.00066793675,"about_ca_system_score_gemma":0.0012282581,"threshold_uncertainty_score":0.104521155},"labels":[],"label_agreement":null},{"id":"W2463732179","doi":"","title":"[Searching for and processing professional and scientific data 2/2].","year":2011,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Obsolescence; Computer science; Value (mathematics); Subject (documents); Data science; Knowledge management; Sociology of scientific knowledge; Business; World Wide Web; Machine learning; Sociology; Marketing","score_opus":0.16049173239094913,"score_gpt":0.3164719878232611,"score_spread":0.15598025543231195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2463732179","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060358294,0.02254473,0.55139315,0.052126605,0.008008538,0.0074075027,0.0999832,0.038314976,0.2141854],"genre_scores_gemma":[0.020963272,0.016039018,0.82124984,0.013380099,0.001991367,0.0035847526,0.07422861,0.0055560484,0.043006983],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9894546,0.0033784232,0.002752292,0.0007183707,0.0033290382,0.00036721816],"domain_scores_gemma":[0.97762376,0.008511848,0.001951125,0.004858332,0.0064067543,0.00064814533],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010852763,0.0009917931,0.0013604971,0.013916749,0.0017674378,0.0064546335,0.0022283578,0.0024021603,0.040800933],"category_scores_gemma":[0.043090507,0.00086021965,0.0020884979,0.016522294,0.0011383243,0.0075762225,0.0065450324,0.0015726179,0.058851637],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000106313426,0.0000348403,0.0010776112,0.005765995,0.00014989507,0.00045828123,0.00089554,0.00014376182,0.002762202,0.023588575,0.49706653,0.4679505],"study_design_scores_gemma":[0.00003596871,0.000018324668,0.001563858,0.0013993067,0.00007995509,0.0007650137,0.00049852417,0.0006707729,0.0012017922,0.019665409,0.9740485,0.000052642117],"about_ca_topic_score_codex":0.0035900222,"about_ca_topic_score_gemma":0.00621496,"teacher_disagreement_score":0.040800933,"about_ca_system_score_codex":0.00077947,"about_ca_system_score_gemma":0.008473902,"threshold_uncertainty_score":0.13649273},"labels":[],"label_agreement":null},{"id":"W2479033895","doi":"10.6082/hzt1q-ygb10","title":"Finding Our Way through Phenotypes","year":2015,"lang":"en","type":"article","venue":"UKnowledge (University of Kentucky)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Phenotype; Biology; Genetics; Gene","score_opus":0.04900077040518896,"score_gpt":0.2681472746544255,"score_spread":0.21914650424923654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2479033895","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09327549,0.007813904,0.56998694,0.112241566,0.0030093298,0.0003462956,0.016368816,0.005525082,0.19143268],"genre_scores_gemma":[0.5288229,0.01196377,0.377129,0.014730217,0.00073283963,0.00048275045,0.012224931,0.004081509,0.04983209],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99624085,0.0014426917,0.0002657232,0.001143577,0.00073260017,0.00017454647],"domain_scores_gemma":[0.9915221,0.0033355136,0.0007875852,0.0024918942,0.001298771,0.0005641728],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006460192,0.00083092664,0.00071418047,0.0026956853,0.0021489575,0.008985428,0.0015878334,0.0015231741,0.018529577],"category_scores_gemma":[0.028858442,0.00046674506,0.0012526792,0.0027974516,0.006755473,0.01548792,0.0052203755,0.0033977183,0.0048799305],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014110867,0.000081162754,0.04699478,0.00069171155,0.00014471725,0.0011272278,0.020095551,0.0015512966,0.0038106886,0.53876394,0.07255444,0.31404337],"study_design_scores_gemma":[0.000010409184,0.000041299125,0.015180475,0.00069933926,0.000100566846,0.0010533141,0.009502354,0.0018508005,0.001979172,0.4696581,0.49982774,0.000096389325],"about_ca_topic_score_codex":0.0054050847,"about_ca_topic_score_gemma":0.005547252,"teacher_disagreement_score":0.018529577,"about_ca_system_score_codex":0.0016779198,"about_ca_system_score_gemma":0.0030312052,"threshold_uncertainty_score":0.06198758},"labels":[],"label_agreement":null},{"id":"W2483810189","doi":"10.18848/1832-3669/cgp/v08i05/56328","title":"Health Care Delivery in Remote and Isolated First Nations Communities in Canada: The Need for a Collaborative Health Informatics System","year":2013,"lang":"en","type":"article","venue":"The International Journal of Technology Knowledge and Society","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Health informatics; Health care delivery; Informatics; Health care; Medicine; Business; Nursing; Political science; Public health","score_opus":0.009442438161453483,"score_gpt":0.26196634211730463,"score_spread":0.25252390395585117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2483810189","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7158629,0.0062107663,0.004616781,0.23046447,0.00030015348,0.00096441316,0.0020886338,0.00030528128,0.03918664],"genre_scores_gemma":[0.97418684,0.0025477544,0.011936872,0.0077158343,0.00006274161,0.0002347032,0.00077457674,0.00004686339,0.0024938453],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99247134,0.0022706313,0.000541095,0.00065848645,0.0019268963,0.0021315883],"domain_scores_gemma":[0.96430314,0.0084568225,0.0032455388,0.00090156775,0.008937962,0.0141549995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007709354,0.00036896247,0.0006128889,0.0038709054,0.019093968,0.010502091,0.0038024632,0.0021328807,0.0037085759],"category_scores_gemma":[0.025230542,0.00048355444,0.00057508505,0.0074306233,0.004115136,0.0045904242,0.009030443,0.0027804251,0.0002573696],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032719248,0.0009460119,0.4961289,0.0011086876,0.00044003647,0.0025434638,0.107697524,0.0044575687,0.001176737,0.031168977,0.056366406,0.29763842],"study_design_scores_gemma":[0.00015134626,0.00018570726,0.46438116,0.0026498728,0.0003052195,0.0007482693,0.41823956,0.010926374,0.0005078572,0.014362836,0.08721952,0.00032231084],"about_ca_topic_score_codex":0.9891306,"about_ca_topic_score_gemma":0.99362767,"teacher_disagreement_score":0.11565134,"about_ca_system_score_codex":0.11565134,"about_ca_system_score_gemma":0.37131366,"threshold_uncertainty_score":0.8391131},"labels":[],"label_agreement":null},{"id":"W2491036503","doi":"10.4018/978-1-4666-1803-9.ch016","title":"Natural Language Processing and Machine Learning Techniques Help Achieve a Better Medical Practice","year":2012,"lang":"en","type":"book-chapter","venue":"Advances in medical technologies and clinical practice book series","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Medical practice; Artificial intelligence; Medical information; Clinical Practice; Data science; Natural language processing; Machine learning; Information retrieval","score_opus":0.01299459924541557,"score_gpt":0.35718151234689327,"score_spread":0.3441869131014777,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2491036503","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029555222,0.16019969,0.33898404,0.064918935,0.005981209,0.0003613871,0.0012813555,0.0026670631,0.42265087],"genre_scores_gemma":[0.021469403,0.14065355,0.51016015,0.014668935,0.0054593054,0.00035933426,0.0020336297,0.0012376223,0.3039581],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99917823,0.00020129088,0.00005735105,0.00013100173,0.00040309294,0.000029048082],"domain_scores_gemma":[0.9982516,0.0013666677,0.00007149778,0.00008639575,0.00018068604,0.000043103613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014927157,0.0010665262,0.00063127023,0.0020445425,0.00069892645,0.0043151653,0.00084206805,0.0016184942,0.024404583],"category_scores_gemma":[0.0030642238,0.00040847747,0.00076548336,0.002242598,0.0018439798,0.0072595538,0.0011436564,0.0028489921,0.018259065],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022524651,0.000086058586,0.00024196199,0.0014439031,0.00005403517,0.00019625406,0.0010383092,0.0014349459,0.003652906,0.2420129,0.24257907,0.5072371],"study_design_scores_gemma":[0.0000071804648,0.000016876871,0.0002790478,0.0005473136,0.000012246387,0.00030860657,0.00024248393,0.0013020644,0.00073490135,0.14096336,0.8555643,0.000021589221],"about_ca_topic_score_codex":0.00061514403,"about_ca_topic_score_gemma":0.0014034951,"teacher_disagreement_score":0.024404583,"about_ca_system_score_codex":0.0010550362,"about_ca_system_score_gemma":0.0015036953,"threshold_uncertainty_score":0.081641495},"labels":[],"label_agreement":null},{"id":"W2498651415","doi":"10.4018/978-1-61520-733-6.ch016","title":"Implications of Web 2.0 Technology on Healthcare","year":2010,"lang":"en","type":"book-chapter","venue":"Advances in healthcare information systems and administration book series","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"World Wide Web; RSS; Folksonomy; Computer science; Metadata; Annotation; Web resource; Bookmarking; Microblogging; Health care; Social media","score_opus":0.0136800020942805,"score_gpt":0.30028288721051183,"score_spread":0.28660288511623133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2498651415","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021291745,0.28783277,0.009191764,0.09049691,0.01337688,0.00005973285,0.00018859167,0.0003534678,0.5963707],"genre_scores_gemma":[0.043810997,0.49144423,0.025334993,0.0719901,0.014764107,0.00018707558,0.00049474294,0.0004360985,0.3515377],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99930024,0.00028008374,0.00003252123,0.000055902914,0.00026648457,0.00006475823],"domain_scores_gemma":[0.998089,0.0015057719,0.000048182606,0.0000637905,0.00019640502,0.00009679955],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013692661,0.0005757265,0.0003270825,0.0015474388,0.000848935,0.0065104533,0.00059179886,0.0026001239,0.014154603],"category_scores_gemma":[0.0021775195,0.00029062098,0.00039689968,0.0028522222,0.0022786246,0.0063229566,0.0018026644,0.0027275977,0.0073152743],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019758088,0.000042939937,0.00019876765,0.0007505313,0.000013725409,0.00033465697,0.0012305032,0.000727152,0.00054336095,0.39377537,0.32367468,0.27868858],"study_design_scores_gemma":[0.0000045223114,0.000008960737,0.00027148714,0.00062581,0.0000040055447,0.00035941275,0.00027814746,0.00026286248,0.00014487463,0.04758259,0.9504447,0.00001262751],"about_ca_topic_score_codex":0.0030600256,"about_ca_topic_score_gemma":0.0035873877,"teacher_disagreement_score":0.014154603,"about_ca_system_score_codex":0.0025768247,"about_ca_system_score_gemma":0.002416889,"threshold_uncertainty_score":0.047351897},"labels":[],"label_agreement":null},{"id":"W2507494804","doi":"10.3233/978-1-61499-678-1-399","title":"An Ontological Model of Behaviour Theory to Generate Personalized Action Plans to Modify Behaviours","year":2016,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Operationalization; Ontology; Action (physics); Correctness; Formative assessment; Knowledge management; Process management; Psychology; Engineering","score_opus":0.09213204651514222,"score_gpt":0.4012603674610217,"score_spread":0.3091283209458795,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2507494804","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050532552,0.00020546312,0.9707587,0.0015265592,0.00009861886,0.000442983,0.0015335794,0.0007993486,0.01958162],"genre_scores_gemma":[0.11549032,0.00058977905,0.8718911,0.000558942,0.00006319985,0.0012748343,0.003599831,0.00019005456,0.006341947],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99800247,0.0007063993,0.0002819456,0.0003736145,0.0005184319,0.00011707587],"domain_scores_gemma":[0.99750745,0.0012566522,0.00019224265,0.0005741599,0.0003740087,0.000095494666],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002581006,0.00082085025,0.0005617579,0.0029029576,0.001246594,0.0032329946,0.0019610333,0.001352814,0.004939206],"category_scores_gemma":[0.0057316977,0.00059019186,0.003262274,0.0017548291,0.0021889305,0.0035062148,0.0017396641,0.0022754343,0.0014632236],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000634732,0.0003008899,0.00283662,0.00042427008,0.00015429879,0.00043334858,0.0023920508,0.055806216,0.0022370697,0.84873706,0.0053798566,0.08123477],"study_design_scores_gemma":[0.00006874815,0.00009053516,0.0021032067,0.0004072919,0.00023974896,0.00037886723,0.0008415704,0.3255987,0.0026435526,0.5449787,0.12257243,0.00007667486],"about_ca_topic_score_codex":0.02223373,"about_ca_topic_score_gemma":0.020470891,"teacher_disagreement_score":0.02223373,"about_ca_system_score_codex":0.0024660397,"about_ca_system_score_gemma":0.0045330534,"threshold_uncertainty_score":0.044208646},"labels":[],"label_agreement":null},{"id":"W2507883388","doi":"10.18653/v1/w16-3005","title":"VERSE: Event and Relation Extraction in the BioNLP 2016 Shared Task","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"","keywords":"Computer science; Relationship extraction; Biomedical text mining; Task (project management); Relation (database); Event (particle physics); Natural language processing; Information retrieval; Data mining; Text mining; Engineering; Physics","score_opus":0.014446756060164715,"score_gpt":0.2827033617687736,"score_spread":0.2682566057086089,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2507883388","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10661884,0.0035974174,0.3006875,0.0041402387,0.0020544436,0.0028355578,0.34820467,0.20375769,0.028103653],"genre_scores_gemma":[0.087037235,0.0004968595,0.22772262,0.00055298797,0.00021415278,0.0015059893,0.66496193,0.0066857394,0.0108224265],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.991321,0.0023417505,0.0011361545,0.002808544,0.001833831,0.0005586717],"domain_scores_gemma":[0.98152345,0.010942825,0.00063164043,0.00354844,0.0024222156,0.00093140674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00880685,0.005100827,0.0022805443,0.0047960524,0.0024318066,0.002963563,0.0034437007,0.0042295023,0.023990322],"category_scores_gemma":[0.025239587,0.001182219,0.0026991193,0.0031448938,0.000878717,0.006900779,0.006707052,0.0035119522,0.018223228],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029967239,0.0012759089,0.0071971016,0.0043869223,0.000572142,0.0022372715,0.00132412,0.014469281,0.036839478,0.0062959855,0.57029897,0.35210615],"study_design_scores_gemma":[0.0013836834,0.00085756095,0.018355219,0.0005246184,0.0004387512,0.0030037565,0.0016877772,0.23287784,0.08436061,0.024853103,0.6311779,0.0004792611],"about_ca_topic_score_codex":0.015852531,"about_ca_topic_score_gemma":0.02137038,"teacher_disagreement_score":0.023990322,"about_ca_system_score_codex":0.0019741743,"about_ca_system_score_gemma":0.0044728047,"threshold_uncertainty_score":0.08025569},"labels":[],"label_agreement":null},{"id":"W2509885322","doi":"10.1093/database/baw121","title":"BioCreative V BioC track overview: collaborative biocurator assistant task for BioGRID","year":2016,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Sinai Hospital; Lunenfeld-Tanenbaum Research Institute; Université de Montréal; Institute for Research in Immunology and Cancer","funders":"National Institute of General Medical Sciences; Biotechnology and Biological Sciences Research Council; National Institutes of Health","keywords":"Computer science; Usability; Annotation; Task (project management); Interoperability; Classifier (UML); World Wide Web; Information retrieval; Data curation; Crowdsourcing; Natural language processing; Artificial intelligence; Human–computer interaction","score_opus":0.025597812753064744,"score_gpt":0.3119697022994703,"score_spread":0.2863718895464056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2509885322","genre_codex":"software","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035829023,0.0025977823,0.3038677,0.0034863944,0.0014643279,0.006897336,0.24771187,0.35443047,0.043715134],"genre_scores_gemma":[0.039651774,0.00068088266,0.2936366,0.0009860748,0.00029604818,0.0059092394,0.5991034,0.030356215,0.029379824],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9935921,0.0015126264,0.00072201947,0.0015559075,0.0019049147,0.0007124021],"domain_scores_gemma":[0.9805921,0.0039999904,0.0011671623,0.006180745,0.005391428,0.0026685416],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017059505,0.0027356187,0.0020838818,0.004400817,0.0033523028,0.0053510265,0.005879368,0.0027017584,0.026871176],"category_scores_gemma":[0.019773025,0.0016887544,0.0020269088,0.005271286,0.00057171716,0.005621374,0.0064990637,0.0022481862,0.035236835],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012464983,0.00049029215,0.004354779,0.0010577758,0.00018219509,0.0002166773,0.0010051377,0.0017593723,0.010960668,0.0021165397,0.87807703,0.098533124],"study_design_scores_gemma":[0.000998131,0.00082263036,0.012419098,0.00024849354,0.00017140096,0.0005523171,0.0007160143,0.02359476,0.027296461,0.00377033,0.92911166,0.00029874418],"about_ca_topic_score_codex":0.032511916,"about_ca_topic_score_gemma":0.031033915,"teacher_disagreement_score":0.032511916,"about_ca_system_score_codex":0.002791755,"about_ca_system_score_gemma":0.007993827,"threshold_uncertainty_score":0.09022039},"labels":[],"label_agreement":null},{"id":"W2513382469","doi":"","title":"Lesser-spotted zebras: Their care and feeding.","year":2016,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Alberta Health Services","funders":"","keywords":"Corollary; Maxim; Computer science; Medicine; Philosophy; Mathematics; Combinatorics; Epistemology","score_opus":0.014856781677599781,"score_gpt":0.2111368015423442,"score_spread":0.19628001986474441,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2513382469","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029953683,0.18237588,0.011030771,0.67150575,0.017788485,0.000072589755,0.0036278928,0.0010068095,0.08263823],"genre_scores_gemma":[0.32374504,0.20056318,0.03753488,0.23394004,0.009926939,0.00018770671,0.0058204867,0.00063223275,0.18764953],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99932885,0.00023784276,0.0000715329,0.00008576044,0.00019726735,0.00007884447],"domain_scores_gemma":[0.9990447,0.00022993312,0.00020302774,0.00006793726,0.0002087743,0.00024559986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012822896,0.00051041116,0.00041995745,0.0016994567,0.00189679,0.0018375954,0.0006747822,0.0019946187,0.012314426],"category_scores_gemma":[0.0045339963,0.00032817823,0.00043197747,0.0013659929,0.0031619382,0.0038830193,0.0025035383,0.0032267098,0.005152178],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016635412,0.000026153499,0.014167258,0.00079440844,0.00007885022,0.0012706675,0.008104558,0.000059730115,0.0021614698,0.01721176,0.68039566,0.2755631],"study_design_scores_gemma":[0.000012905939,0.00004584507,0.021207863,0.00094512827,0.00003528128,0.005937417,0.0096676145,0.00008886985,0.0006063588,0.010867353,0.9505329,0.00005253948],"about_ca_topic_score_codex":0.014246494,"about_ca_topic_score_gemma":0.055602737,"teacher_disagreement_score":0.014246494,"about_ca_system_score_codex":0.0013158695,"about_ca_system_score_gemma":0.0018507333,"threshold_uncertainty_score":0.04119593},"labels":[],"label_agreement":null},{"id":"W2516876363","doi":"10.1126/scisignal.aah4406","title":"Quantitative human cell encyclopedia","year":2016,"lang":"en","type":"article","venue":"Science Signaling","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute of General Medical Sciences; National Institutes of Health; Icahn School of Medicine at Mount Sinai","keywords":"Encyclopedia; Computational biology; Cataloging; Biology; Human cell; Cell biology; Computer science; World Wide Web; Library science; Biochemistry; Gene","score_opus":0.022254692420408274,"score_gpt":0.30946907835064463,"score_spread":0.28721438593023635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2516876363","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001026904,0.0051225713,0.009014801,0.0002684481,0.00030495162,0.00013591746,0.9650612,0.006765984,0.012299266],"genre_scores_gemma":[0.00446868,0.0073181265,0.0163314,0.0004177375,0.00010569424,0.0004538285,0.96183145,0.0013423779,0.0077307355],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992267,0.00010292203,0.00020890676,0.00016319478,0.00023870416,0.000059528316],"domain_scores_gemma":[0.9973074,0.0011683884,0.00018653522,0.00045947675,0.00073939597,0.00013883541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011007031,0.0014981941,0.001165697,0.010573412,0.00061574555,0.0027652634,0.0012888421,0.000683805,0.077257104],"category_scores_gemma":[0.00449052,0.00046078642,0.00073118176,0.015503066,0.0003061482,0.0013203865,0.0014906292,0.00087762595,0.03987276],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031079687,0.000048017562,0.001807496,0.008265764,0.00017511677,0.0003549547,0.00020636884,0.0008721204,0.007714182,0.01199363,0.7930058,0.17524572],"study_design_scores_gemma":[0.000030312729,0.000018798459,0.0026743084,0.00044778164,0.0000852317,0.0003182743,0.000058061367,0.00046532782,0.0027945961,0.0025074882,0.99056834,0.000031570187],"about_ca_topic_score_codex":0.007999164,"about_ca_topic_score_gemma":0.009016264,"teacher_disagreement_score":0.077257104,"about_ca_system_score_codex":0.00089097855,"about_ca_system_score_gemma":0.0035015328,"threshold_uncertainty_score":0.25845075},"labels":[],"label_agreement":null},{"id":"W2520572539","doi":"10.1101/073460","title":"Semi-Automated Identification of Ontological Labels in the Biomedical Literature with goldi","year":2016,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Ontario; University of Toronto; University of Ottawa; Centre for Addiction and Mental Health","funders":"","keywords":"Computer science; Leverage (statistics); Identification (biology); Parsing; Scope (computer science); Ontology; Data science; Function (biology); Dependency grammar; Artificial intelligence; Information retrieval; Natural language processing; Epistemology; Programming language","score_opus":0.010252606456181142,"score_gpt":0.24175487163466333,"score_spread":0.23150226517848219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2520572539","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041271474,0.002382473,0.862084,0.0031580105,0.0003220627,0.0009830145,0.020357292,0.05861675,0.010824953],"genre_scores_gemma":[0.045009356,0.00029763093,0.9372693,0.00023875597,0.00007866266,0.0003419757,0.01422972,0.0008399735,0.0016945979],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99362916,0.0017291928,0.00084918295,0.0015082138,0.001968393,0.00031588337],"domain_scores_gemma":[0.9822916,0.010285641,0.0017840489,0.002215965,0.0029968654,0.00042594047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008120967,0.001295127,0.0013315241,0.020475185,0.0019637968,0.004955871,0.0018053459,0.001477993,0.005298978],"category_scores_gemma":[0.032706168,0.0007555118,0.0017767358,0.008566969,0.0012666851,0.0032179079,0.0060688523,0.001877715,0.0054541687],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049337815,0.0002655012,0.016465437,0.0030883446,0.00034579294,0.00061245024,0.0029422548,0.007857165,0.028261874,0.04547291,0.1155903,0.7786047],"study_design_scores_gemma":[0.00018232457,0.00022211882,0.017669171,0.0013458877,0.0003571666,0.0014853791,0.003317263,0.49832088,0.063136995,0.17726694,0.2363807,0.00031523328],"about_ca_topic_score_codex":0.00469479,"about_ca_topic_score_gemma":0.0100159105,"teacher_disagreement_score":0.020475185,"about_ca_system_score_codex":0.0020090262,"about_ca_system_score_gemma":0.004303107,"threshold_uncertainty_score":0.042948306},"labels":[],"label_agreement":null},{"id":"W2522746188","doi":"10.3389/fmed.2016.00039","title":"Distributed Cognition and Process Management Enabling Individualized Translational Research: The NIH Undiagnosed Diseases Program Experience","year":2016,"lang":"en","type":"article","venue":"Frontiers in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Hospital for Sick Children","funders":"National Human Genome Research Institute; Research Institute for Information Technology, Kyushu University; Common Fund; National Institutes of Health","keywords":"Translational research; Multidisciplinary approach; Process (computing); Usability; Scalability; Computer science; Cloud computing; Knowledge management; Cognition; Process management; Data science; Medicine; Engineering; Human–computer interaction; Psychology; Neuroscience","score_opus":0.04959252115637274,"score_gpt":0.37715817804481205,"score_spread":0.3275656568884393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2522746188","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47456262,0.0078092525,0.24478024,0.184304,0.0012731336,0.0015690965,0.00067484676,0.004099315,0.0809274],"genre_scores_gemma":[0.70967877,0.0038591,0.26385024,0.010639182,0.0006562894,0.0008539139,0.00091359654,0.00042769476,0.009121214],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.96612376,0.025266325,0.001143953,0.0016485859,0.004354436,0.0014630086],"domain_scores_gemma":[0.9351922,0.041102428,0.0017863642,0.0067632347,0.005869862,0.0092858095],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06340382,0.0004849309,0.00046167275,0.0010681129,0.0029390403,0.008285167,0.0031929626,0.0019626776,0.0020993834],"category_scores_gemma":[0.040544577,0.00035759038,0.0006011391,0.0016163912,0.004285438,0.0064893174,0.009364454,0.0040549235,0.00045278826],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010602545,0.0057594706,0.0289047,0.0009887072,0.00020074125,0.0009957444,0.071836166,0.0046074158,0.008252149,0.06793602,0.062378123,0.74708056],"study_design_scores_gemma":[0.0013608056,0.005268182,0.028202958,0.0010818503,0.00042190595,0.001640343,0.05847061,0.04105771,0.020946337,0.17201291,0.66917,0.00036641458],"about_ca_topic_score_codex":0.0068130344,"about_ca_topic_score_gemma":0.006548379,"teacher_disagreement_score":0.06340382,"about_ca_system_score_codex":0.0050363,"about_ca_system_score_gemma":0.013755511,"threshold_uncertainty_score":0.33531553},"labels":[],"label_agreement":null},{"id":"W2523769980","doi":"10.18162/ritpu.2007.138","title":"Un multi-outil adapté au parcours cognitif de l’étudiant en traduction spécialisée : application à la biomédecine","year":2007,"lang":"fr","type":"article","venue":"Revue internationale des technologies en pédagogie universitaire","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Humanities; Philosophy","score_opus":0.023033788906123186,"score_gpt":0.2856663779589817,"score_spread":0.2626325890528585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2523769980","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0856103,0.002675563,0.8427987,0.0043926896,0.00066332286,0.0013779397,0.00054694735,0.015425616,0.046508927],"genre_scores_gemma":[0.31607586,0.0014803726,0.64591616,0.0017234163,0.0003341175,0.0011437786,0.0007813211,0.0010581265,0.031486843],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9957211,0.0016178748,0.0002766986,0.0010199021,0.0010895615,0.00027482983],"domain_scores_gemma":[0.9868868,0.006161337,0.0011775332,0.0017352987,0.0026718653,0.0013671393],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0065259426,0.0016431833,0.0010391442,0.0022739582,0.0010752789,0.007120157,0.0025032011,0.002835172,0.016406557],"category_scores_gemma":[0.014593484,0.00088141084,0.000878265,0.0009829603,0.0015100893,0.0066687106,0.0057153455,0.0016044757,0.0043397914],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020500736,0.0016099109,0.013619341,0.0018637831,0.0002280185,0.000849943,0.009344957,0.007468162,0.05204736,0.023305764,0.01440692,0.87320584],"study_design_scores_gemma":[0.00064060144,0.0062477686,0.04063426,0.0029920726,0.0011715482,0.004908206,0.009974161,0.3357892,0.10054198,0.079789415,0.4162533,0.0010574834],"about_ca_topic_score_codex":0.0031288376,"about_ca_topic_score_gemma":0.0039850078,"teacher_disagreement_score":0.016406557,"about_ca_system_score_codex":0.0012304373,"about_ca_system_score_gemma":0.0025942347,"threshold_uncertainty_score":0.054885447},"labels":[],"label_agreement":null},{"id":"W2524580445","doi":"","title":"Something Old, Something New: Identifying Knowledge Source in Bio-Events","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Computer science; Annotation; Event (particle physics); Domain knowledge; Context (archaeology); Knowledge extraction; Sentence; Natural language processing; Domain (mathematical analysis); Bridge (graph theory); Information retrieval; Knowledge-based systems; Information extraction; Artificial intelligence; Data science","score_opus":0.029191963855970438,"score_gpt":0.3055341479342584,"score_spread":0.276342184078288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2524580445","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29375842,0.008634893,0.6127175,0.0055061933,0.0006070943,0.0015341268,0.026101213,0.024218021,0.026922524],"genre_scores_gemma":[0.38396248,0.0024896932,0.57703465,0.00043383194,0.0002485543,0.00042895923,0.027613524,0.00092385727,0.006864549],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99734914,0.0005816572,0.0005271208,0.0007707812,0.0006377389,0.00013346513],"domain_scores_gemma":[0.9806012,0.01347085,0.0023917397,0.0013208189,0.0018293484,0.00038607788],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037381933,0.0009413147,0.00075964816,0.017334923,0.0019025384,0.0044805994,0.0015495546,0.001960854,0.0036854392],"category_scores_gemma":[0.017992225,0.0004623091,0.00073215406,0.008932478,0.001144862,0.007603674,0.002616083,0.0012961298,0.0028110913],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011282056,0.0003727292,0.059410516,0.004295708,0.00017755946,0.0049085086,0.019784411,0.0026585793,0.056609116,0.020607172,0.04091612,0.78913146],"study_design_scores_gemma":[0.00027389303,0.00024766152,0.14819404,0.002233254,0.0008446949,0.011043607,0.024806786,0.17774305,0.17421608,0.0835425,0.37638524,0.00046924723],"about_ca_topic_score_codex":0.0046456233,"about_ca_topic_score_gemma":0.0052459086,"teacher_disagreement_score":0.017334923,"about_ca_system_score_codex":0.0013963286,"about_ca_system_score_gemma":0.0013707882,"threshold_uncertainty_score":0.019769728},"labels":[],"label_agreement":null},{"id":"W2536039441","doi":"10.3233/978-1-61499-678-1-322","title":"Transcription of Case Report Forms from Unstructured Referral Letters: A Semantic Text Analytics Approach","year":2016,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University","funders":"","keywords":"Referral; SNOMED CT; Ontology; Computer science; Natural language processing; Analytics; Information retrieval; Artificial intelligence; Data science; Medicine; Family medicine; Terminology; Linguistics","score_opus":0.04595609939916981,"score_gpt":0.33064528190703624,"score_spread":0.28468918250786646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2536039441","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031560086,0.00061785954,0.9272762,0.0025025532,0.00038665585,0.002579705,0.01320268,0.011272848,0.010601455],"genre_scores_gemma":[0.087911285,0.000568033,0.8895868,0.00033950093,0.0002059081,0.0008363076,0.016927436,0.00079546217,0.002829265],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99240714,0.003079677,0.0012645589,0.0010981146,0.0019200343,0.00023042136],"domain_scores_gemma":[0.9680559,0.017735608,0.0035688295,0.0033376415,0.006806042,0.00049601094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0058052726,0.0015299887,0.0008684869,0.0107266335,0.0014501335,0.0041825995,0.0015823141,0.0012135948,0.006225598],"category_scores_gemma":[0.024518946,0.0005951242,0.0009663528,0.005796011,0.0015061143,0.00351291,0.0030480293,0.0018317228,0.005921131],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007797953,0.00043996843,0.0099735875,0.003527851,0.00013351328,0.0055124625,0.019628353,0.010750799,0.08381348,0.02481286,0.0464336,0.79419374],"study_design_scores_gemma":[0.00021721765,0.0004292671,0.022730205,0.0027773383,0.00032375075,0.008793128,0.050382648,0.16960289,0.1314233,0.13651618,0.47614434,0.00065978296],"about_ca_topic_score_codex":0.0028138838,"about_ca_topic_score_gemma":0.0033484816,"teacher_disagreement_score":0.0107266335,"about_ca_system_score_codex":0.0013915402,"about_ca_system_score_gemma":0.0031974532,"threshold_uncertainty_score":0.030701578},"labels":[],"label_agreement":null},{"id":"W2543755080","doi":"10.1177/0193945916673815","title":"Explorative Analyses of Nursing Research Data","year":2016,"lang":"en","type":"article","venue":"Western Journal of Nursing Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cancer Care Ontario","funders":"National Institutes of Health; National Institute for Health and Care Research; Patient-Centered Outcomes Research Institute","keywords":"Metadata; Standardization; Scope (computer science); Computer science; Metadata repository; Meta Data Services; Big data; Data element; Nursing; Data science; World Wide Web; Medicine; Data mining","score_opus":0.6548863625811201,"score_gpt":0.6113940098487308,"score_spread":0.04349235273238927,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2543755080","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84057564,0.0025818278,0.111771084,0.0046898094,0.00024377709,0.0037265997,0.022193242,0.000549548,0.013668471],"genre_scores_gemma":[0.84237325,0.0007705603,0.13986813,0.00069541595,0.000115233765,0.0048924657,0.009547741,0.00015373743,0.0015834939],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9496486,0.032375988,0.0061095054,0.0026801678,0.0075775636,0.0016080976],"domain_scores_gemma":[0.7458587,0.20455953,0.016739648,0.01812928,0.013674372,0.0010385495],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04858485,0.0008646008,0.000893899,0.026146295,0.0021117206,0.0050338563,0.001426256,0.0007761757,0.0021004705],"category_scores_gemma":[0.14096512,0.0004283399,0.0015955596,0.023166081,0.0022610098,0.003578692,0.0058824127,0.0010698476,0.00036471052],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092318427,0.0006001222,0.4029517,0.0071976604,0.0015830684,0.0066187438,0.25604934,0.003398903,0.019950956,0.069710344,0.010941641,0.22007434],"study_design_scores_gemma":[0.00010888878,0.00062395155,0.39148676,0.003950749,0.00104148,0.0039966065,0.345433,0.018092329,0.019114215,0.08238176,0.1333975,0.00037269326],"about_ca_topic_score_codex":0.003821048,"about_ca_topic_score_gemma":0.00593074,"teacher_disagreement_score":0.04858485,"about_ca_system_score_codex":0.0025141647,"about_ca_system_score_gemma":0.0052924454,"threshold_uncertainty_score":0.2569443},"labels":[],"label_agreement":null},{"id":"W2545390300","doi":"","title":"What is a transcript","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science","score_opus":0.015251493809675973,"score_gpt":0.2750156097513121,"score_spread":0.2597641159416361,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2545390300","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0670972,0.0033155044,0.35045972,0.055433113,0.0085796835,0.0011569903,0.2213692,0.013360216,0.27922836],"genre_scores_gemma":[0.3962359,0.0063029244,0.23792875,0.006108695,0.0043402235,0.0010442027,0.20889685,0.0055610994,0.13358136],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.997544,0.00047513557,0.00027268613,0.0007315737,0.00075774395,0.00021886929],"domain_scores_gemma":[0.99535525,0.0021152445,0.00038639127,0.00085026596,0.0010483459,0.0002446119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014777157,0.00066563475,0.0007062521,0.0027658853,0.0017182283,0.0033981525,0.00072375796,0.0011000512,0.031098397],"category_scores_gemma":[0.008052959,0.00036117647,0.0013452165,0.002932702,0.001361778,0.006601091,0.001453986,0.0018941397,0.014773715],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007974984,0.00016628293,0.02162827,0.0021235452,0.00025793852,0.0042821527,0.0056502665,0.001967798,0.032344975,0.26128072,0.29945013,0.37005037],"study_design_scores_gemma":[0.00004511541,0.00006961744,0.011824123,0.00061261625,0.00024553106,0.0024678994,0.004272676,0.004636064,0.013487028,0.15608758,0.80612224,0.00012952255],"about_ca_topic_score_codex":0.0049171555,"about_ca_topic_score_gemma":0.0039253025,"teacher_disagreement_score":0.031098397,"about_ca_system_score_codex":0.0011171417,"about_ca_system_score_gemma":0.00287981,"threshold_uncertainty_score":0.10403448},"labels":[],"label_agreement":null},{"id":"W2546705495","doi":"","title":"WaterlooClarke: TREC 2015 Clinical Decision Support Track","year":2015,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Information retrieval; Mean reciprocal rank; Search engine; Reciprocal; Clinical decision support system; Test (biology); Rank (graph theory); Decision support system; Data mining","score_opus":0.08993383767107038,"score_gpt":0.3806690461193441,"score_spread":0.2907352084482737,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2546705495","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036413785,0.029699324,0.029374493,0.05483222,0.00578544,0.008391572,0.737739,0.031545933,0.0662182],"genre_scores_gemma":[0.03816479,0.0046054334,0.044603344,0.004586429,0.0008665811,0.0021267696,0.8756398,0.001120256,0.028286653],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99337137,0.0020278313,0.000941037,0.00081918243,0.0023713582,0.00046917534],"domain_scores_gemma":[0.9685702,0.009968861,0.0018621355,0.002199364,0.014742782,0.002656629],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015881011,0.0026891783,0.0023911097,0.0074915034,0.0023047873,0.004924522,0.003775383,0.0032494783,0.032339767],"category_scores_gemma":[0.033751845,0.0009227378,0.0012474112,0.0057223383,0.001085069,0.0045769764,0.0022177494,0.002943837,0.019928355],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030709655,0.00021051537,0.0009338949,0.0009810907,0.00010314845,0.00007845258,0.00005550222,0.00088606233,0.0013194153,0.0003740012,0.96546936,0.029281504],"study_design_scores_gemma":[0.0031988139,0.0010601465,0.026358034,0.0015543945,0.0005518914,0.0007390126,0.0005435813,0.048678894,0.01784947,0.006426466,0.8926049,0.0004343838],"about_ca_topic_score_codex":0.16449954,"about_ca_topic_score_gemma":0.27007198,"teacher_disagreement_score":0.16449954,"about_ca_system_score_codex":0.009992594,"about_ca_system_score_gemma":0.016980717,"threshold_uncertainty_score":0.32708406},"labels":[],"label_agreement":null},{"id":"W2547991767","doi":"10.1111/1365-2745.12698","title":"Towards a thesaurus of plant characteristics: an ecological contribution","year":2016,"lang":"en","type":"article","venue":"Journal of Ecology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":154,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"Deutsches Zentrum für integrative Biodiversitätsforschung Halle-Jena-Leipzig; Centre National de la Recherche Scientifique; Consejo Nacional de Investigaciones Científicas y Técnicas; Deutsche Forschungsgemeinschaft; Fondo para la Investigación Científica y Tecnológica; Inter-American Institute for Global Change Research; National Science Foundation","keywords":"Terminology; Ontology; Computer science; Thesaurus; Context (archaeology); Information retrieval; Semantics (computer science); Quality (philosophy); Resource (disambiguation); Semantic Web; World Wide Web; Data science; Ecology; Geography; Artificial intelligence; Linguistics; Biology","score_opus":0.014804135974905689,"score_gpt":0.2692910647887432,"score_spread":0.2544869288138375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2547991767","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026289247,0.0040433053,0.917242,0.006730935,0.0017567378,0.0007731328,0.0064045954,0.0021504122,0.03460962],"genre_scores_gemma":[0.07290898,0.0041371607,0.8982965,0.0011142332,0.0005740421,0.0010097144,0.013538301,0.0010825394,0.007338522],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99706024,0.00096717064,0.00084663584,0.00040353023,0.0006150531,0.00010728666],"domain_scores_gemma":[0.9900568,0.003668153,0.0007530305,0.0020856601,0.0026987027,0.00073765893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067616506,0.0004830707,0.00088660244,0.013046314,0.0017359119,0.0060245898,0.002131735,0.0019087701,0.005570205],"category_scores_gemma":[0.013802425,0.00056863116,0.0015538121,0.0107250335,0.00286752,0.0104168635,0.0048931893,0.00280502,0.002049324],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000114520255,0.0001989409,0.005210982,0.0030015928,0.00013264062,0.00085550855,0.010186262,0.0051194853,0.008862653,0.5431556,0.05358266,0.3695791],"study_design_scores_gemma":[0.00002738816,0.00008414654,0.004989413,0.002079853,0.00011359171,0.0014038809,0.0029713425,0.01225312,0.0019430377,0.14763345,0.8263903,0.000110383655],"about_ca_topic_score_codex":0.0042367345,"about_ca_topic_score_gemma":0.0031665342,"teacher_disagreement_score":0.013046314,"about_ca_system_score_codex":0.0020955568,"about_ca_system_score_gemma":0.0051648174,"threshold_uncertainty_score":0.03575945},"labels":[],"label_agreement":null},{"id":"W2548614243","doi":"10.1109/itab.2010.5687670","title":"Detection and normalization of medical terms using domain-specific term frequency and adaptive ranking","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Unified Medical Language System; Computer science; Information retrieval; Normalization (sociology); Ranking (information retrieval); Term (time); Thesaurus; Sentence; Natural language processing; Artificial intelligence","score_opus":0.015447592321949715,"score_gpt":0.26359450736472134,"score_spread":0.24814691504277162,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2548614243","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10409809,0.0017517101,0.8851212,0.00029302086,0.00015270026,0.0005851155,0.0009437155,0.0054242113,0.0016302405],"genre_scores_gemma":[0.1834122,0.00051569985,0.81159234,0.00009799871,0.00014562852,0.0004026261,0.002254225,0.00030610204,0.0012730281],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9950623,0.0012685772,0.0007859621,0.00087880314,0.0018096628,0.00019465748],"domain_scores_gemma":[0.9904625,0.0040455665,0.0012364154,0.0011169129,0.0029402555,0.00019841244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003827444,0.0009322855,0.0018265975,0.01206866,0.0007436699,0.0017572416,0.0015241092,0.00095039577,0.0011615775],"category_scores_gemma":[0.017918946,0.00035239125,0.0012935029,0.0078032454,0.0006116705,0.002480497,0.0008933559,0.00081239763,0.0011563072],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038979243,0.00038434332,0.011255433,0.00056833256,0.00024774447,0.00020620873,0.00034391056,0.0076518073,0.08544695,0.0024287098,0.0039046681,0.8871721],"study_design_scores_gemma":[0.00042636367,0.0010418665,0.078714,0.00015775606,0.0010639852,0.0039882986,0.00080844667,0.6845594,0.18750732,0.014441153,0.026792016,0.00049938704],"about_ca_topic_score_codex":0.0045765634,"about_ca_topic_score_gemma":0.005230459,"teacher_disagreement_score":0.01206866,"about_ca_system_score_codex":0.00093221717,"about_ca_system_score_gemma":0.0017286112,"threshold_uncertainty_score":0.020241678},"labels":[],"label_agreement":null},{"id":"W2550915265","doi":"10.1186/s12859-016-1352-7","title":"Introducing Explorer of Taxon Concepts with a case study on spider measurement matrix building","year":2016,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"Consejo Nacional de Investigaciones Científicas y Técnicas; University of Florida; Stony Brook University; National Science Foundation","keywords":"Spider; Taxon; Data science; Biology; Evolutionary biology; Computer science; Matrix (chemical analysis); Computational biology; Geography; Ecology; Chemistry","score_opus":0.051289187817629374,"score_gpt":0.3172159791570347,"score_spread":0.26592679133940533,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2550915265","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3638966,0.001960649,0.52545184,0.00667228,0.0005452874,0.0015826811,0.011991289,0.01458602,0.07331341],"genre_scores_gemma":[0.2673993,0.0007257533,0.7007939,0.00052164093,0.00012323688,0.00071131706,0.007515323,0.0029726003,0.019236796],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974015,0.0012654212,0.0001778416,0.00039285855,0.00066678907,0.00009559295],"domain_scores_gemma":[0.98943627,0.008036718,0.00047917763,0.0008607982,0.00082186045,0.00036508692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003084648,0.0005733395,0.00024924582,0.0034588838,0.00145146,0.0025706945,0.0010911692,0.0010711842,0.011540936],"category_scores_gemma":[0.011169539,0.00030861745,0.00081582257,0.0030961507,0.0013400327,0.0038830047,0.0027471974,0.0011095221,0.0025051008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048385427,0.0006092006,0.034860924,0.0038356031,0.000098281555,0.028494177,0.14730005,0.0074599907,0.053825334,0.0883555,0.10378802,0.5308891],"study_design_scores_gemma":[0.000055191776,0.00018117634,0.015520187,0.0009018101,0.000058003163,0.010442158,0.03129274,0.027542131,0.024727508,0.026954576,0.8621727,0.0001519314],"about_ca_topic_score_codex":0.0028579482,"about_ca_topic_score_gemma":0.007798171,"teacher_disagreement_score":0.011540936,"about_ca_system_score_codex":0.0010163002,"about_ca_system_score_gemma":0.0011054013,"threshold_uncertainty_score":0.038608313},"labels":[],"label_agreement":null},{"id":"W2557385283","doi":"10.1093/nar/gkw1039","title":"The Human Phenotype Ontology in 2017","year":2016,"lang":"en","type":"review","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":801,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"SickKids Foundation; University of Toronto; Children's Hospital of Eastern Ontario; Hospital for Sick Children; University of Ottawa","funders":"Basic Energy Sciences; U.S. National Library of Medicine; Office of Science; National Institute for Health and Care Research; Bundesministerium für Bildung und Forschung; U.S. Department of Energy; European Commission; National Institutes of Health; NIHR Biomedical Research Centre, Royal Marsden NHS Foundation Trust/Institute of Cancer Research; Deutsche Forschungsgemeinschaft; Moorfields Eye Hospital NHS Foundation Trust; European Science Foundation; Cold Spring Harbor Laboratory; British Heart Foundation; National Human Genome Research Institute; Wellcome Trust; E-Rare","keywords":"Phenotype; Biology; Computational biology; Ontology; Clinical phenotype; Terminology; Pipeline (software); Disease; Bioinformatics; Controlled vocabulary; Data science; Genetics; Computer science; Gene; Information retrieval; Pathology; Medicine","score_opus":0.1460962601911709,"score_gpt":0.46240248005805834,"score_spread":0.31630621986688745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2557385283","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005074436,0.9433863,0.009125201,0.011577943,0.004485541,0.00011834372,0.0026857813,0.0003818268,0.02773168],"genre_scores_gemma":[0.0039946926,0.9550678,0.009055067,0.0075288424,0.0011792701,0.00017669528,0.0070484322,0.0001462566,0.015803007],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991666,0.00017336907,0.00016509311,0.00011674522,0.00031905703,0.00005911335],"domain_scores_gemma":[0.9990625,0.00038799483,0.00009442294,0.00006112109,0.00031662837,0.00007727668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021028025,0.00089178246,0.0009414223,0.0039053077,0.0004812596,0.001918034,0.0012819909,0.0018193187,0.006951622],"category_scores_gemma":[0.0036629774,0.00040554247,0.00077783986,0.0038629766,0.0013494131,0.003360971,0.0024194564,0.003198896,0.0068310327],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005948258,0.00003472574,0.00024808102,0.0057990383,0.00005405161,0.00032465285,0.00019709772,0.00038104333,0.0011955836,0.04297575,0.17690074,0.7718297],"study_design_scores_gemma":[0.000003506798,0.00000471948,0.00018750054,0.00093605235,0.000011409793,0.00030335068,0.000022184011,0.000031745294,0.00013840212,0.002503246,0.9958508,0.000007216209],"about_ca_topic_score_codex":0.007693619,"about_ca_topic_score_gemma":0.0066252835,"teacher_disagreement_score":0.007693619,"about_ca_system_score_codex":0.0027484584,"about_ca_system_score_gemma":0.0069865854,"threshold_uncertainty_score":0.023255527},"labels":[],"label_agreement":null},{"id":"W2559958825","doi":"","title":"[Future Perspective of Pharmacoepidemiology in the \"Big Data Era\" and the Growth of Information Sources].","year":2016,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"","keywords":"Pharmacoepidemiology; Enthusiasm; Observational study; Big data; Confusion; Data science; Perspective (graphical); Medicine; Computer science; Risk analysis (engineering); Data mining; Psychology; Pharmacology; Medical prescription; Pathology","score_opus":0.03400357232454273,"score_gpt":0.27112127128852104,"score_spread":0.2371176989639783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2559958825","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004172903,0.44189155,0.009306807,0.52192724,0.014358882,0.00007261944,0.00115159,0.00029407075,0.010580018],"genre_scores_gemma":[0.015049772,0.75812566,0.03860121,0.15018858,0.029385045,0.00026279682,0.0020159837,0.00011697381,0.0062540197],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9923963,0.004898864,0.00069186423,0.00040505998,0.0012802029,0.00032773893],"domain_scores_gemma":[0.94934165,0.03301078,0.0027925526,0.0023240955,0.010168454,0.0023623407],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021563308,0.00083198986,0.0015023084,0.0068503437,0.0011878443,0.0077551366,0.0024988628,0.008011432,0.011071252],"category_scores_gemma":[0.030531734,0.0003384681,0.0018273653,0.013458876,0.004869009,0.01885927,0.002983155,0.0048549334,0.0038650616],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013742397,0.00006845556,0.0017183118,0.008680461,0.0002705467,0.00033932136,0.00054599,0.0006027026,0.00035282882,0.14155519,0.52200675,0.32372206],"study_design_scores_gemma":[0.000036862668,0.00004383357,0.00216649,0.00801825,0.0001245762,0.0005150679,0.0007835039,0.00055414555,0.00013220734,0.1562132,0.83133936,0.0000725064],"about_ca_topic_score_codex":0.0071271425,"about_ca_topic_score_gemma":0.008630116,"teacher_disagreement_score":0.021563308,"about_ca_system_score_codex":0.0040079895,"about_ca_system_score_gemma":0.013115018,"threshold_uncertainty_score":0.114039004},"labels":[],"label_agreement":null},{"id":"W2561809471","doi":"10.1002/pra2.2016.14505301116","title":"Library of congress subject heading (LCSH) browsing and natural language searching","year":2016,"lang":"en","type":"article","venue":"Proceedings of the Association for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Hierarchy; Subject (documents); Natural language; USable; Controlled vocabulary; Vocabulary; Heading (navigation); Subject access","score_opus":0.0051767441317691265,"score_gpt":0.24643698860169533,"score_spread":0.2412602444699262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2561809471","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21099177,0.002441191,0.4815964,0.004139404,0.00027698273,0.0020166207,0.015483694,0.14191523,0.14113879],"genre_scores_gemma":[0.3835734,0.0011300724,0.5813186,0.0006742014,0.00012146643,0.00055328367,0.01258047,0.00415347,0.015895048],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99825186,0.0006130138,0.00025751546,0.00023792451,0.0005300697,0.000109676555],"domain_scores_gemma":[0.98790586,0.0069047944,0.00089695374,0.0020751962,0.0015291711,0.00068805186],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0033164378,0.0004293046,0.0005769193,0.0068425043,0.0010694136,0.0041196723,0.0010779743,0.0006035935,0.010705427],"category_scores_gemma":[0.015903413,0.00038723776,0.0004778516,0.005569217,0.00073475525,0.005446399,0.0032129774,0.00051240565,0.003334938],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00095411274,0.00028876256,0.02222307,0.0023199415,0.00013842828,0.00080798386,0.012217284,0.0026543476,0.03254171,0.050552864,0.10772663,0.76757485],"study_design_scores_gemma":[0.00033414704,0.00044772055,0.029590156,0.0012055058,0.00027197044,0.0012318172,0.008277384,0.10446857,0.04616436,0.052376572,0.7552953,0.0003364824],"about_ca_topic_score_codex":0.00929699,"about_ca_topic_score_gemma":0.016510088,"teacher_disagreement_score":0.9958803,"about_ca_system_score_codex":0.001005847,"about_ca_system_score_gemma":0.0020205546,"threshold_uncertainty_score":0.035813272},"labels":[],"label_agreement":null},{"id":"W2562842110","doi":"10.1186/s13326-016-0099-4","title":"An ontological analysis of medical Bayesian indicators of performance","year":2017,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Intégré Universitaire de Santé et de Services Sociaux du Saguenay–Lac-Saint-Jean; Centre Intégré Universitaire de Santé et de Services Sociaux du Centre-Sud-de-l'Île-de-Montréal; Centre Hospitalier Universitaire de Sherbrooke; Université de Sherbrooke","funders":"","keywords":"Ontology; Computer science; Context (archaeology); Probabilistic logic; Bayesian probability; Test (biology); Sensitivity (control systems); Open Biomedical Ontologies; Value (mathematics); Representation (politics); Artificial intelligence; Natural language processing; Information retrieval; Machine learning; Upper ontology; Ontology alignment; Semantic Web; Epistemology","score_opus":0.015757272201087275,"score_gpt":0.3293877132986925,"score_spread":0.3136304410976052,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2562842110","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07000079,0.0015483762,0.8730212,0.00644527,0.00017742038,0.0004269265,0.0029328507,0.00051951455,0.044927705],"genre_scores_gemma":[0.63160187,0.001237779,0.35940805,0.0006475653,0.00030438646,0.000464793,0.003089059,0.0001175729,0.0031289656],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9935999,0.0018887657,0.00061514083,0.00082479924,0.002662194,0.00040930777],"domain_scores_gemma":[0.98492855,0.007897505,0.0017504736,0.0013197211,0.0035346567,0.0005689723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008196646,0.0006257032,0.0005480353,0.008954164,0.001624042,0.0053809783,0.0014343864,0.0011693353,0.0029231093],"category_scores_gemma":[0.033051427,0.00039876276,0.0017628656,0.007240506,0.0029755547,0.005787677,0.002838774,0.0013631735,0.00052702863],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000095972006,0.00012087052,0.013486969,0.00024394614,0.00012182608,0.00041819346,0.0020828473,0.015155212,0.0010434124,0.8925943,0.0035404284,0.07109608],"study_design_scores_gemma":[0.000032046562,0.000072791045,0.015428018,0.00042489002,0.00028868968,0.0007075469,0.0018524501,0.18023111,0.00149606,0.7582337,0.041142754,0.00008989472],"about_ca_topic_score_codex":0.013649516,"about_ca_topic_score_gemma":0.006134166,"teacher_disagreement_score":0.013649516,"about_ca_system_score_codex":0.005848641,"about_ca_system_score_gemma":0.0035652167,"threshold_uncertainty_score":0.04334855},"labels":[],"label_agreement":null},{"id":"W2563349699","doi":"10.18653/v1/w16-6113","title":"Hybrid methods for ICD-10 coding of death certificates","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Nautical Research Society","funders":"European Commission","keywords":"ICD-10; Coding (social sciences); Computer science; Statistics; Medicine; Mathematics","score_opus":0.05916374058954748,"score_gpt":0.37753889367765847,"score_spread":0.318375153088111,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2563349699","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027186787,0.00094576407,0.96237755,0.0005467164,0.00020182542,0.0003007907,0.0017718149,0.0036770443,0.0029917245],"genre_scores_gemma":[0.24285199,0.00066018355,0.73952967,0.0002527128,0.00022001134,0.00058078096,0.009577207,0.00036624505,0.0059611565],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964217,0.0017554639,0.0003044444,0.0006149407,0.0007780928,0.00012530173],"domain_scores_gemma":[0.9933149,0.004201227,0.00029130786,0.0008994478,0.0011679215,0.00012522901],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048686992,0.00082128646,0.00070778944,0.0041434923,0.0006578263,0.0015366098,0.0012772682,0.0008735437,0.0056547206],"category_scores_gemma":[0.012622114,0.00028279066,0.00097628153,0.00331049,0.00042858516,0.0014481236,0.0022722613,0.0014092079,0.004763794],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030648848,0.00018495921,0.008571677,0.0002589795,0.00027093416,0.00009055533,0.0003564903,0.02228553,0.002535347,0.008476816,0.013824253,0.9428379],"study_design_scores_gemma":[0.00013853058,0.00014247658,0.009661419,0.00014983071,0.000106676154,0.00056207954,0.000743073,0.90937966,0.007382818,0.048855864,0.022796338,0.00008125112],"about_ca_topic_score_codex":0.0041493205,"about_ca_topic_score_gemma":0.0066534774,"teacher_disagreement_score":0.0056547206,"about_ca_system_score_codex":0.0006013548,"about_ca_system_score_gemma":0.0017142474,"threshold_uncertainty_score":0.025748432},"labels":[],"label_agreement":null},{"id":"W2573117738","doi":"10.1093/database/baw147","title":"The BioC-BioGRID corpus: full text articles annotated for curation of protein–protein and genetic interactions","year":2016,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Sinai Hospital; Lunenfeld-Tanenbaum Research Institute; Université de Montréal; Institute for Research in Immunology and Cancer","funders":"Biotechnology and Biological Sciences Research Council; National Institutes of Health","keywords":"Data curation; Annotation; Crowdsourcing; Computer science; Task (project management); Information retrieval; Gene Annotation; Pipeline (software); World Wide Web; Data science; Natural language processing; Computational biology; Artificial intelligence; Biology; Gene; Genome; Genetics; Engineering","score_opus":0.019403511412564904,"score_gpt":0.2759297630851905,"score_spread":0.2565262516726256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2573117738","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02363438,0.013649165,0.007220298,0.0024760917,0.00094611605,0.0010179456,0.9286564,0.0030116583,0.01938801],"genre_scores_gemma":[0.02132499,0.004333302,0.015191063,0.00056925457,0.0002523643,0.0017302745,0.95038193,0.0009158172,0.005300917],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962913,0.00077279407,0.00084486057,0.0007200956,0.001172252,0.00019869814],"domain_scores_gemma":[0.9774449,0.0147281755,0.001787689,0.0014998424,0.003728505,0.0008109412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030985258,0.0016414218,0.0013614661,0.021265294,0.0027849877,0.003135238,0.0019892638,0.002031925,0.02024647],"category_scores_gemma":[0.016202081,0.0006109745,0.0007237725,0.026096946,0.0016896743,0.0020637321,0.0036159789,0.001786069,0.014526825],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048636826,0.00013312469,0.0026490018,0.020831965,0.00015256519,0.002653658,0.0022938934,0.0007540624,0.014003166,0.0039262385,0.8940946,0.058021266],"study_design_scores_gemma":[0.00014066405,0.000033112345,0.012213608,0.0016603625,0.00013472003,0.0009985316,0.0010919318,0.00073396176,0.004950676,0.0012088843,0.9767725,0.000060958933],"about_ca_topic_score_codex":0.010150199,"about_ca_topic_score_gemma":0.028258102,"teacher_disagreement_score":0.021265294,"about_ca_system_score_codex":0.0024817395,"about_ca_system_score_gemma":0.005512308,"threshold_uncertainty_score":0.0677312},"labels":[],"label_agreement":null},{"id":"W2574627809","doi":"10.3233/978-1-61499-660-6-285","title":"A Molecular Structure Ontology for Medicinal Chemistry","year":2016,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Ontology; Chemistry; Computer science; Computational biology; Information retrieval; Epistemology; Philosophy; Biology","score_opus":0.021940616144699727,"score_gpt":0.28363175734362295,"score_spread":0.2616911411989232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2574627809","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001239679,0.019876432,0.60944015,0.013835234,0.0034112188,0.00056564924,0.005262053,0.0030016678,0.34336793],"genre_scores_gemma":[0.018347535,0.025298193,0.80271226,0.006018395,0.0010629336,0.0007489395,0.009183698,0.0010844108,0.13554361],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994185,0.00009943956,0.00007407283,0.00010043657,0.00027521,0.000032339965],"domain_scores_gemma":[0.9994135,0.00030463777,0.000036856967,0.000100524805,0.000102302,0.00004221862],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011417287,0.0010318436,0.0005704902,0.0034651158,0.0013736986,0.0032940158,0.0016441788,0.001859378,0.020891935],"category_scores_gemma":[0.0018267644,0.00076018577,0.0014147607,0.0045596077,0.0019792647,0.009219069,0.0024340937,0.004139661,0.009348279],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007741051,0.000028151555,0.000051614363,0.00049665605,0.000011728603,0.00011120403,0.0003296605,0.001240556,0.0011461905,0.78842103,0.09437126,0.11378422],"study_design_scores_gemma":[0.0000069549737,0.0000050327326,0.00007052305,0.00021508161,0.000006442115,0.00024299206,0.000040476396,0.001600765,0.00032457718,0.15247093,0.845005,0.000011281341],"about_ca_topic_score_codex":0.0035127548,"about_ca_topic_score_gemma":0.0056287064,"teacher_disagreement_score":0.020891935,"about_ca_system_score_codex":0.00337369,"about_ca_system_score_gemma":0.0031344425,"threshold_uncertainty_score":0.0698905},"labels":[],"label_agreement":null},{"id":"W2579634769","doi":"","title":"Modeling Life Science Knowledge with OWL 1.1.","year":2008,"lang":"en","type":"article","venue":"Research Publications (Maastricht University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Ontology; Computer science; Knowledge representation and reasoning; Web Ontology Language; Process ontology; Open Biomedical Ontologies; Representation (politics); Upper ontology; OWL-S; Data science; Knowledge management; Suggested Upper Merged Ontology; Semantic Web; Information retrieval; Domain knowledge; Artificial intelligence; Epistemology","score_opus":0.13148983573196216,"score_gpt":0.34375124501192855,"score_spread":0.2122614092799664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2579634769","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005098646,0.0009946496,0.9587806,0.0029515566,0.00025972037,0.00053887104,0.0059015146,0.00512041,0.020354003],"genre_scores_gemma":[0.07405775,0.001999488,0.8939303,0.0016583969,0.00018177087,0.0007947175,0.017319322,0.00094308273,0.009115161],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99738616,0.00064962666,0.0004608911,0.00025452874,0.0010586972,0.00019023367],"domain_scores_gemma":[0.99729556,0.0010325285,0.00049094285,0.0005269442,0.0005117445,0.00014235807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039123683,0.0008092448,0.000735034,0.0029461838,0.0012738118,0.0047819233,0.002350228,0.0014546057,0.0032772352],"category_scores_gemma":[0.007455518,0.0008108014,0.0017772474,0.0035055433,0.0012287198,0.0064566215,0.0029407763,0.0024790766,0.0019782966],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000101656224,0.000272256,0.002180124,0.0016448187,0.00026449747,0.0016883428,0.0027264263,0.030447625,0.007857059,0.7057213,0.056630224,0.19046561],"study_design_scores_gemma":[0.000053702784,0.000038928276,0.0009054294,0.00054605294,0.00013823679,0.00081489387,0.00050086086,0.09159635,0.0054512103,0.4544273,0.4454511,0.00007585421],"about_ca_topic_score_codex":0.01577369,"about_ca_topic_score_gemma":0.01563708,"teacher_disagreement_score":0.01577369,"about_ca_system_score_codex":0.001867853,"about_ca_system_score_gemma":0.003766006,"threshold_uncertainty_score":0.031363726},"labels":[],"label_agreement":null},{"id":"W2582792602","doi":"","title":"Detecting semantic changes in Alzheimer’s disease with vector space models","year":2016,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; Disease; Space (punctuation); Medicine; Pathology","score_opus":0.0302507815983099,"score_gpt":0.2936973496375037,"score_spread":0.2634465680391938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2582792602","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7153819,0.0032037639,0.26839274,0.0012062953,0.00016869669,0.00029479826,0.0054834667,0.003967141,0.0019011438],"genre_scores_gemma":[0.9340327,0.0005317718,0.059358787,0.0001326485,0.000044480184,0.00012044726,0.0048758793,0.00008048608,0.0008228886],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983979,0.00060839986,0.00018869713,0.0003024856,0.00036901986,0.00013345151],"domain_scores_gemma":[0.9959372,0.003069145,0.00021443273,0.00018298041,0.00049958867,0.00009658597],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030420132,0.00093708653,0.00084036146,0.004642628,0.00046324386,0.0015813796,0.00081289164,0.000952925,0.0009940042],"category_scores_gemma":[0.0072942,0.00017251167,0.0012886757,0.0024171579,0.00040526132,0.001989576,0.0008062003,0.00073341385,0.00035958356],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028041478,0.001222556,0.0711288,0.0005906096,0.0010547368,0.00057423953,0.00035398948,0.21148048,0.008699782,0.005788444,0.009770092,0.686532],"study_design_scores_gemma":[0.000046574227,0.00022150579,0.005412978,0.000031142124,0.00017280345,0.00017062208,0.00017421243,0.9824923,0.0034757112,0.007036544,0.0007468592,0.000018684179],"about_ca_topic_score_codex":0.020351091,"about_ca_topic_score_gemma":0.012166571,"teacher_disagreement_score":0.020351091,"about_ca_system_score_codex":0.0011366252,"about_ca_system_score_gemma":0.001154074,"threshold_uncertainty_score":0.040465236},"labels":[],"label_agreement":null},{"id":"W2584889691","doi":"10.2196/medinform.6918","title":"Ontology-Driven Search and Triage: Design of a Web-Based Visual Interface for MEDLINE","year":2017,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Triage; Interface (matter); MEDLINE; World Wide Web; Ontology; Information retrieval; Medicine; Medical emergency","score_opus":0.04417655343339854,"score_gpt":0.385797005319146,"score_spread":0.34162045188574747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2584889691","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028740564,0.00025725446,0.9083785,0.00073737715,0.00009996527,0.0017330605,0.0016219297,0.054564252,0.0038672108],"genre_scores_gemma":[0.0923355,0.0002595518,0.8943239,0.00055072294,0.000041257837,0.0023951393,0.0018639637,0.0031149334,0.005114993],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998214,0.00061596686,0.00031606635,0.00032001812,0.00041890042,0.0001151483],"domain_scores_gemma":[0.98871267,0.008748934,0.0004768799,0.00050359574,0.00091814477,0.0006398083],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.003922177,0.0015883059,0.00083102,0.0028089187,0.00045044094,0.0029738573,0.0025297706,0.0018331786,0.011649958],"category_scores_gemma":[0.016821736,0.0009880089,0.0010993603,0.0010564105,0.000691276,0.0031860706,0.0026962329,0.0012320101,0.0029756464],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004780726,0.0013865023,0.009568755,0.0062523372,0.00048460255,0.0061790654,0.01938195,0.02262273,0.20373262,0.023055904,0.065663904,0.63689095],"study_design_scores_gemma":[0.0023563632,0.0019033883,0.009453449,0.0019055641,0.0005513736,0.0057426533,0.0031281768,0.52371275,0.111030266,0.028761638,0.31063095,0.0008233764],"about_ca_topic_score_codex":0.0013570376,"about_ca_topic_score_gemma":0.0013749818,"teacher_disagreement_score":0.99702615,"about_ca_system_score_codex":0.0007017165,"about_ca_system_score_gemma":0.0012979389,"threshold_uncertainty_score":0.038973033},"labels":[],"label_agreement":null},{"id":"W2585015406","doi":"10.4018/ijitwe.2017040102","title":"Semantic Reconciliation of Electronic Health Records Using Semantic Web Technologies","year":2017,"lang":"en","type":"article","venue":"International Journal of Information Technology and Web Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Computer science; RDF; Semantic reasoner; Semantic Web; Information retrieval; SPARQL; Semantics (computer science); Coding (social sciences); Semantic analytics; Health records; Process (computing); Electronic medical record; World Wide Web; Medical record; Semantic Web Stack; Health care; Artificial intelligence; Programming language; Internet privacy; Medicine","score_opus":0.0075791838438117225,"score_gpt":0.2633674179288234,"score_spread":0.25578823408501167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2585015406","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010550304,0.0011550789,0.9744573,0.0031112067,0.00045429944,0.00045185705,0.0008720806,0.003261353,0.0056865527],"genre_scores_gemma":[0.17113945,0.001784427,0.8136211,0.0014363057,0.00036865947,0.00034479474,0.006861269,0.00070905074,0.0037349171],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97008914,0.0138823325,0.003468565,0.0028230997,0.009008182,0.00072866905],"domain_scores_gemma":[0.9671971,0.010248902,0.0029001108,0.013871273,0.005435393,0.00034717808],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024376914,0.001144698,0.0016951753,0.011309253,0.0025897548,0.008786718,0.004063961,0.0025257007,0.0019881772],"category_scores_gemma":[0.033377293,0.0008300087,0.0034836975,0.010396685,0.0030924971,0.01667277,0.00977863,0.0035693725,0.0008933209],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050728524,0.0004359469,0.0046884487,0.0016054491,0.0010773836,0.0020681466,0.0068646236,0.024774393,0.006893683,0.422841,0.020643,0.50760067],"study_design_scores_gemma":[0.00015277801,0.0001707699,0.0024688523,0.001584211,0.00074616756,0.0015639702,0.0046201563,0.117121965,0.04064774,0.5679098,0.2627655,0.00024816062],"about_ca_topic_score_codex":0.002973008,"about_ca_topic_score_gemma":0.0031721655,"teacher_disagreement_score":0.024376914,"about_ca_system_score_codex":0.0021594225,"about_ca_system_score_gemma":0.0056025693,"threshold_uncertainty_score":0.128919},"labels":[],"label_agreement":null},{"id":"W2587099845","doi":"10.1186/s13040-017-0123-y","title":"Semantics-based plausible reasoning to extend the knowledge coverage of medical knowledge bases for improved clinical decision support","year":2017,"lang":"en","type":"article","venue":"BioData Mining","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Model-based reasoning; Semantics (computer science); Semantic Web; Data science; Artificial intelligence; Information retrieval; Knowledge representation and reasoning","score_opus":0.076125745526062,"score_gpt":0.4277109479327042,"score_spread":0.3515852024066422,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2587099845","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025323633,0.0013276209,0.96383095,0.002332294,0.00007594311,0.00050194375,0.0017298489,0.0021641143,0.0027136982],"genre_scores_gemma":[0.29605663,0.0008023266,0.6975414,0.0007907345,0.00011196602,0.0003284332,0.0038247544,0.0001389423,0.00040477168],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.984385,0.006653038,0.002535575,0.0020387631,0.004039368,0.0003483242],"domain_scores_gemma":[0.95116884,0.037522227,0.0023324566,0.0049690194,0.0034649605,0.00054251315],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.016976617,0.001452508,0.0016331233,0.008250455,0.0011471983,0.004963693,0.0031732935,0.0018535331,0.0031014432],"category_scores_gemma":[0.074995145,0.0008898011,0.0042814417,0.0052235117,0.0018738466,0.009463459,0.005553761,0.0026840128,0.000719651],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010876284,0.0010531265,0.0148470355,0.0027050667,0.0013373173,0.0031093082,0.0037388974,0.31359756,0.012715942,0.12758754,0.009137814,0.5090828],"study_design_scores_gemma":[0.0001967267,0.00015438178,0.0014684091,0.00039494794,0.00041054067,0.0007503456,0.00042858583,0.7074946,0.008411471,0.26583326,0.014343642,0.00011299068],"about_ca_topic_score_codex":0.0049434076,"about_ca_topic_score_gemma":0.006416947,"teacher_disagreement_score":0.016976617,"about_ca_system_score_codex":0.0017239585,"about_ca_system_score_gemma":0.0036418682,"threshold_uncertainty_score":0.08978194},"labels":[],"label_agreement":null},{"id":"W2588463081","doi":"10.1109/ghtc.2016.7857300","title":"Teaching bilingual workshops on data mining in Peru","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Thompson Rivers University","funders":"","keywords":"Terminology; Fluency; Experiential learning; Computer science; Mathematics education; Neuroscience of multilingualism; Biomedicine; Artificial intelligence; Psychology; Linguistics","score_opus":0.05138660967687258,"score_gpt":0.33418711788844957,"score_spread":0.282800508211577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2588463081","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8950532,0.0011041657,0.027891958,0.012767959,0.00023058668,0.0005858981,0.00036527845,0.00057019456,0.061430752],"genre_scores_gemma":[0.93514574,0.0017657223,0.033497397,0.00289898,0.00012941648,0.00097338244,0.00071061327,0.00012114743,0.024757622],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978387,0.0013877464,0.00007258537,0.00020187133,0.00015718606,0.0003420282],"domain_scores_gemma":[0.99593145,0.0019928487,0.0002551379,0.00023203246,0.0004644962,0.0011240985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003660843,0.0005455566,0.00033626272,0.0009310876,0.0040627876,0.002104285,0.0010531172,0.0008789666,0.012157564],"category_scores_gemma":[0.008483823,0.0003541089,0.0004808987,0.0012028562,0.0013427307,0.0020999478,0.008838987,0.0014076212,0.0021087795],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006201997,0.0064141285,0.031768784,0.0014073043,0.000056429202,0.007051371,0.5635457,0.0009095087,0.01305346,0.015895205,0.023941508,0.3353364],"study_design_scores_gemma":[0.0003895213,0.0019006279,0.039573655,0.00071173394,0.000051307874,0.0032248474,0.42238495,0.0034808246,0.008107695,0.016801747,0.5032289,0.00014425041],"about_ca_topic_score_codex":0.0023103259,"about_ca_topic_score_gemma":0.0059496337,"teacher_disagreement_score":0.012157564,"about_ca_system_score_codex":0.0016477753,"about_ca_system_score_gemma":0.004117024,"threshold_uncertainty_score":0.04067111},"labels":[],"label_agreement":null},{"id":"W2591194248","doi":"10.29173/cais753","title":"Traditional versus Blogosphere Information Landscapes: The Case of Diabetes and HbA1c","year":2013,"lang":"fr","type":"article","venue":"Proceedings of the Annual Conference of CAIS / Actes du congrès annuel de l ACSI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Blogosphere; Library science; Humanities; Computer science; Philosophy; World Wide Web; The Internet","score_opus":0.019902004523093517,"score_gpt":0.23018924164282342,"score_spread":0.2102872371197299,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591194248","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.961284,0.0017911944,0.0056417235,0.0036847158,0.00006940209,0.000041722487,0.00083069806,0.00010150295,0.026555037],"genre_scores_gemma":[0.9962535,0.00034173275,0.0021147975,0.00009047438,0.000036931433,0.000020239591,0.00024261662,0.00003041822,0.0008691974],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986363,0.000750314,0.000056763824,0.00013841056,0.00025886917,0.00015928721],"domain_scores_gemma":[0.97778016,0.018465394,0.0011372226,0.0006132987,0.0013827815,0.0006210476],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024704211,0.00023090397,0.00033573402,0.00481684,0.0014383005,0.0053074453,0.00037009048,0.00090236583,0.002743876],"category_scores_gemma":[0.01709851,0.00017752731,0.00033641327,0.0059534987,0.0022174309,0.0070528844,0.0016876832,0.0005381597,0.0002761487],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025699052,0.00052929565,0.4377272,0.0018158967,0.00039410754,0.006141691,0.143079,0.008387164,0.009445204,0.121122964,0.017711695,0.25107586],"study_design_scores_gemma":[0.00013230425,0.00035356413,0.5947382,0.0004712734,0.0002977277,0.003334251,0.19274482,0.036214314,0.0024343333,0.09685254,0.072229855,0.00019681659],"about_ca_topic_score_codex":0.006423661,"about_ca_topic_score_gemma":0.007841745,"teacher_disagreement_score":0.006423661,"about_ca_system_score_codex":0.0011568922,"about_ca_system_score_gemma":0.000401669,"threshold_uncertainty_score":0.0130649805},"labels":[],"label_agreement":null},{"id":"W2591383972","doi":"10.29173/cais600","title":"Preliminary Observations on Health Query Terms in a University OPAC: Transaction Log and Co-occurrence Analyses","year":2013,"lang":"fr","type":"article","venue":"Proceedings of the Annual Conference of CAIS / Actes du congrès annuel de l ACSI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Humanities; Computer science; Art","score_opus":0.05691541716665287,"score_gpt":0.3006175442880264,"score_spread":0.2437021271213735,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2591383972","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98722595,0.00034107393,0.005523765,0.00041510732,0.000023978286,0.00013797355,0.0041768895,0.00030097275,0.0018543171],"genre_scores_gemma":[0.98269135,0.0002336781,0.00910061,0.000080701015,0.000056580582,0.00014508102,0.0060603456,0.00007857579,0.0015531537],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9921257,0.0019387997,0.0010763337,0.0007285693,0.003605746,0.0005248514],"domain_scores_gemma":[0.9269912,0.05098279,0.0061006187,0.0034720157,0.010452153,0.0020011333],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0045431154,0.0003484164,0.0005800635,0.008492283,0.0010963425,0.0023020448,0.00073621765,0.00065837544,0.0019261222],"category_scores_gemma":[0.031199804,0.00022954129,0.0004838043,0.0111721875,0.0006764002,0.002042444,0.0009938007,0.0009123977,0.00080814987],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023550577,0.00082264445,0.79472643,0.00073719193,0.00016521302,0.0023809841,0.016597304,0.0037709647,0.039368175,0.0011298888,0.004072205,0.13387385],"study_design_scores_gemma":[0.000048644415,0.0009752404,0.9087091,0.000058358266,0.00012307818,0.0024265966,0.014813294,0.036113076,0.02391082,0.00094233674,0.0117425,0.00013693525],"about_ca_topic_score_codex":0.016098041,"about_ca_topic_score_gemma":0.015724855,"teacher_disagreement_score":0.9915077,"about_ca_system_score_codex":0.0010270928,"about_ca_system_score_gemma":0.0015656324,"threshold_uncertainty_score":0.032008708},"labels":[],"label_agreement":null},{"id":"W2592222179","doi":"10.1200/jco.2013.31.15_suppl.6639","title":"Do clinical trial acronyms affect patients’ interest in clinical trials? A randomized survey.","year":2013,"lang":"en","type":"article","venue":"Journal of Clinical Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Health Sciences Centre; Sunnybrook Health Science Centre","funders":"","keywords":"Medicine; Acronym; Clinical trial; Likert scale; Randomized controlled trial; Respondent; Family medicine; Internal medicine; Psychology","score_opus":0.4403399898659247,"score_gpt":0.5692918533877555,"score_spread":0.12895186352183075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2592222179","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9861681,0.0020441837,0.0010640897,0.0036015634,0.00015080233,0.0031995063,0.0011616922,0.000034903238,0.0025751132],"genre_scores_gemma":[0.9944674,0.0005851526,0.001107477,0.0011011144,0.0000819274,0.0023280845,0.00016287336,0.0000049989,0.00016093529],"study_design_codex":"observational","study_design_gemma":"randomized_trial","domain_scores_codex":[0.8985945,0.08177852,0.011461875,0.0024203763,0.0040926044,0.0016520028],"domain_scores_gemma":[0.7585199,0.15092894,0.07141199,0.005286156,0.005076791,0.008776206],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.057183307,0.00032117433,0.0009688262,0.0013924814,0.0007125071,0.0015277112,0.0005242894,0.0016780228,0.0055974447],"category_scores_gemma":[0.1486566,0.000678995,0.0015946878,0.002021775,0.0017521098,0.0030608634,0.0011792302,0.0020389906,0.0010485577],"study_design_candidate":"randomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.03073269,0.0085122725,0.85856044,0.0033478762,0.0012513411,0.00016762888,0.0065244385,0.00049530144,0.00054885086,0.000724687,0.006985983,0.08214851],"study_design_scores_gemma":[0.0122613935,0.05866598,0.8973931,0.0012493959,0.00085920765,0.0004218298,0.00811432,0.0033869145,0.0007567757,0.00089981133,0.015734673,0.00025663935],"about_ca_topic_score_codex":0.0005250581,"about_ca_topic_score_gemma":0.00059025537,"teacher_disagreement_score":0.9428167,"about_ca_system_score_codex":0.0016142228,"about_ca_system_score_gemma":0.0025928887,"threshold_uncertainty_score":0.30241787},"labels":[],"label_agreement":null},{"id":"W2592348612","doi":"10.1016/j.cmpb.2017.03.003","title":"Leveraging medical taxonomies to improve knowledge management within online communities of practice: The knowledge maps system","year":2017,"lang":"en","type":"article","venue":"Computer Methods and Programs in Biomedicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Lexicon; World Wide Web; Online community; Knowledge translation; Field (mathematics); Data science; Knowledge management; Artificial intelligence","score_opus":0.07779712056365762,"score_gpt":0.41097329796975984,"score_spread":0.3331761774061022,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2592348612","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21350239,0.0042897565,0.641654,0.012134135,0.00083414104,0.002538773,0.031719267,0.06790678,0.025420742],"genre_scores_gemma":[0.20759058,0.0014286016,0.7637896,0.00065268565,0.00015729484,0.00066163944,0.021110192,0.0011443384,0.0034650785],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9953002,0.0016832537,0.00065144233,0.0009289664,0.0012704484,0.00016574583],"domain_scores_gemma":[0.976233,0.013533752,0.0023419624,0.00414539,0.0023602,0.0013857406],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070884675,0.00087013195,0.00091741915,0.015517569,0.0019303086,0.005524964,0.0015831859,0.0015930019,0.003668765],"category_scores_gemma":[0.034731228,0.00065213407,0.001134838,0.010870386,0.00085087534,0.012184085,0.008471401,0.0013528018,0.0018197654],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007912217,0.0014439247,0.042406034,0.002371581,0.0008556149,0.00070459204,0.008558227,0.005734951,0.00799651,0.025492819,0.04727721,0.85636735],"study_design_scores_gemma":[0.0004901614,0.00069721084,0.044838388,0.0017028563,0.0012047332,0.0014156406,0.013463371,0.28908724,0.023150204,0.2568003,0.36662367,0.0005262131],"about_ca_topic_score_codex":0.010146731,"about_ca_topic_score_gemma":0.019393064,"teacher_disagreement_score":0.015517569,"about_ca_system_score_codex":0.0011435366,"about_ca_system_score_gemma":0.0044692773,"threshold_uncertainty_score":0.037487864},"labels":[],"label_agreement":null},{"id":"W2595381146","doi":"10.15171/ijhpm.2017.15","title":"Defining Integrated Knowledge Translation and Moving Forward: A Response to Recent Commentaries","year":2017,"lang":"en","type":"article","venue":"International Journal of Health Policy and Management","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":387,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Western University","funders":"","keywords":"Knowledge translation; Translation (biology); Computer science; Data science; Knowledge management; Biology","score_opus":0.05534345557430839,"score_gpt":0.4141588096681419,"score_spread":0.35881535409383347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2595381146","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00018868546,0.00037143336,0.0005743701,0.9947826,0.0035316027,0.000009234291,0.000027486567,0.000023287625,0.000491344],"genre_scores_gemma":[0.0074396273,0.0010455936,0.0042459653,0.97213745,0.012718081,0.000131483,0.00012991314,0.00017200309,0.0019798689],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8550437,0.045629665,0.024102861,0.015182331,0.04812424,0.011917214],"domain_scores_gemma":[0.31163782,0.53402454,0.024164658,0.017174184,0.09561987,0.01737894],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.18047723,0.0018339254,0.0027087904,0.004298127,0.015624363,0.02940874,0.011290193,0.0885489,0.010763314],"category_scores_gemma":[0.42972445,0.002040432,0.0036158192,0.0070542027,0.029879933,0.04756287,0.029331746,0.11999991,0.0055595073],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008800763,0.00007002302,0.0006770147,0.0005162674,0.000053975495,0.0004666468,0.010042731,0.00034692138,0.00058216805,0.07241959,0.89880323,0.01593323],"study_design_scores_gemma":[0.00010306771,0.0000485849,0.0010876447,0.002460627,0.000081745035,0.00041116946,0.012671091,0.0012586527,0.0009581617,0.0901564,0.89050305,0.00025976216],"about_ca_topic_score_codex":0.021217639,"about_ca_topic_score_gemma":0.020712301,"teacher_disagreement_score":0.81952274,"about_ca_system_score_codex":0.018427277,"about_ca_system_score_gemma":0.06259441,"threshold_uncertainty_score":0.9544662},"labels":[],"label_agreement":null},{"id":"W2596493367","doi":"10.1136/bmjopen-2016-015415.209","title":"209: UNLOCKING EVIDENCE REVERSAL IN THE LITERATURE: A KEY TO TERMINOLOGY","year":2017,"lang":"en","type":"article","venue":"BMJ Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Medicine; Terminology; Key (lock); Family medicine; Linguistics","score_opus":0.14788250655707993,"score_gpt":0.4556263056895478,"score_spread":0.30774379913246785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2596493367","genre_codex":"review","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022697167,0.7558465,0.06491072,0.14011754,0.022310425,0.0024124193,0.0005375051,0.00028174708,0.011313416],"genre_scores_gemma":[0.06827891,0.6057205,0.18041484,0.1116926,0.018047519,0.0113514755,0.0010648994,0.00044254318,0.0029866695],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.72646827,0.17758974,0.061079163,0.00798731,0.02467456,0.0022009492],"domain_scores_gemma":[0.5207246,0.40013126,0.036250968,0.013558932,0.0267123,0.0026218656],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14304699,0.0026716753,0.0068705683,0.043458253,0.0040850523,0.020447508,0.008724269,0.010974913,0.0032312293],"category_scores_gemma":[0.32892928,0.0019236276,0.0044400315,0.038986593,0.04055749,0.034207024,0.015980778,0.013055846,0.001997263],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038863617,0.00007037761,0.0016108314,0.20387287,0.00073722797,0.00094991573,0.031278793,0.00077244174,0.0016695644,0.47646323,0.09442082,0.18776521],"study_design_scores_gemma":[0.00016073977,0.00015004858,0.001526564,0.32951868,0.00088272674,0.002570289,0.010670004,0.000943443,0.0004796621,0.21501587,0.43787026,0.00021175119],"about_ca_topic_score_codex":0.0024887316,"about_ca_topic_score_gemma":0.002617259,"teacher_disagreement_score":0.856953,"about_ca_system_score_codex":0.012669595,"about_ca_system_score_gemma":0.031162351,"threshold_uncertainty_score":0.75651383},"labels":[],"label_agreement":null},{"id":"W2605572910","doi":"10.23889/ijpds.v1i1.257","title":"IMECCHI-DATANETWORK: empowering knowledge generation through international data network","year":2017,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Alberta Health Services; Alberta Children's Hospital; University of Calgary","funders":"","keywords":"Raw data; Computer science; Identifier; Data science; Table (database); Software; Data mining; Matching (statistics); Protocol (science); Observational study; Medicine","score_opus":0.20714538692188547,"score_gpt":0.4843355718825411,"score_spread":0.2771901849606556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2605572910","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006396229,0.0010345752,0.7686174,0.020967036,0.0013776165,0.0035337568,0.061137933,0.088235594,0.04869987],"genre_scores_gemma":[0.045336988,0.001149773,0.82540554,0.0023031072,0.00045061714,0.007952553,0.09559055,0.013081101,0.0087297745],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.95435643,0.030070363,0.0042474302,0.0039934516,0.0060190065,0.0013132571],"domain_scores_gemma":[0.81707627,0.11223994,0.0060408395,0.041970424,0.011535814,0.011136689],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08609819,0.0015713434,0.0015393911,0.008479905,0.002231721,0.011928238,0.005980398,0.0026909183,0.05759726],"category_scores_gemma":[0.19152115,0.0018173652,0.0020015123,0.011493924,0.0022981525,0.01849584,0.028906401,0.005237291,0.023780858],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012785752,0.00050539745,0.0066322805,0.0029563596,0.0003874274,0.0006260049,0.00791734,0.006961683,0.0028717832,0.1347118,0.48029822,0.35485318],"study_design_scores_gemma":[0.00058853737,0.0001412336,0.0022407682,0.0016207879,0.00009178133,0.00020050403,0.0020148815,0.031338163,0.003041401,0.1447488,0.8137769,0.00019624771],"about_ca_topic_score_codex":0.0034468884,"about_ca_topic_score_gemma":0.0034216715,"teacher_disagreement_score":0.08609819,"about_ca_system_score_codex":0.0023332238,"about_ca_system_score_gemma":0.009291568,"threshold_uncertainty_score":0.45533615},"labels":[],"label_agreement":null},{"id":"W2606247321","doi":"10.71781/9974","title":"Application d'algorithmes de bio-informatique à la recherche de patrons de conception","year":2006,"lang":"fr","type":"dissertation","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Art; Library science; Computer science","score_opus":0.08135750366551014,"score_gpt":0.3946744782628927,"score_spread":0.31331697459738256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2606247321","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009288616,0.00061521155,0.9829207,0.00083623175,0.0001524696,0.00015030347,0.00041224898,0.0025153551,0.0031088628],"genre_scores_gemma":[0.08229008,0.00042277342,0.9095235,0.00023305322,0.00010805008,0.00034539707,0.0009529178,0.00038378203,0.0057403906],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99473757,0.0019468492,0.0003923936,0.0012504456,0.0014120464,0.00026072425],"domain_scores_gemma":[0.9760342,0.018212132,0.00064120186,0.0015035882,0.0033221748,0.0002866492],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00859919,0.0013310183,0.0016181098,0.0062831463,0.0021847258,0.008626335,0.0030285174,0.0028358325,0.009357098],"category_scores_gemma":[0.049271207,0.00083688303,0.002184995,0.0046625705,0.002156848,0.005320471,0.0032444193,0.0027572552,0.0027601006],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047038344,0.000243802,0.005843795,0.00047238145,0.0003492972,0.00016596545,0.0016934434,0.08452469,0.0031939785,0.1666169,0.015068486,0.72135687],"study_design_scores_gemma":[0.00015229352,0.000058466172,0.001852683,0.00019475329,0.0001296037,0.00018686298,0.0005573088,0.8582073,0.006009351,0.1044379,0.0281685,0.00004493252],"about_ca_topic_score_codex":0.046414886,"about_ca_topic_score_gemma":0.040670652,"teacher_disagreement_score":0.046414886,"about_ca_system_score_codex":0.0049244915,"about_ca_system_score_gemma":0.0060560755,"threshold_uncertainty_score":0.09228945},"labels":[],"label_agreement":null},{"id":"W2607390459","doi":"10.1093/bioinformatics/btx213","title":"BioCIDER: a Contextualisation InDEx for biological Resources discovery","year":2017,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"Biotechnology and Biological Sciences Research Council; Directorate for Biological Sciences; European Commission","keywords":"Contextualization; Documentation; Computer science; World Wide Web; Source code; Index (typography); Code (set theory); Open source; Software; Open source software; Data science; Programming language","score_opus":0.046255479885533336,"score_gpt":0.30591285235299864,"score_spread":0.2596573724674653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2607390459","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011267122,0.00991121,0.21817979,0.0042079925,0.00095034484,0.0017533978,0.52129143,0.20651232,0.025926407],"genre_scores_gemma":[0.022955503,0.0044038203,0.34487003,0.0013993502,0.0002702727,0.0019361643,0.60733736,0.013172071,0.0036553754],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9958995,0.00074185577,0.00092152326,0.0009410179,0.0012918301,0.00020432695],"domain_scores_gemma":[0.99271303,0.0028446012,0.0009070405,0.0016215757,0.0010611379,0.00085257995],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004445867,0.0025234555,0.0018139577,0.016971184,0.0016846218,0.004349957,0.002077227,0.0017579978,0.017920205],"category_scores_gemma":[0.01929558,0.0011674522,0.0020598157,0.017449023,0.0007752111,0.0068325764,0.008079696,0.0023003893,0.015421417],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012436857,0.0002065261,0.009852496,0.008305646,0.0005845839,0.00078225404,0.0012228529,0.0035144296,0.011293523,0.037547134,0.69615453,0.22929224],"study_design_scores_gemma":[0.00022770418,0.00009149533,0.0077537196,0.0010886556,0.00027501854,0.0005089843,0.0003713861,0.0081389155,0.00631328,0.036950745,0.93810993,0.00017017683],"about_ca_topic_score_codex":0.0073403628,"about_ca_topic_score_gemma":0.012302431,"teacher_disagreement_score":0.017920205,"about_ca_system_score_codex":0.002009777,"about_ca_system_score_gemma":0.0034298599,"threshold_uncertainty_score":0.05994904},"labels":[],"label_agreement":null},{"id":"W2610063871","doi":"10.1186/1471-2105-13-180","title":"Erratum to: A linear classifier based on entity recognition tools and a statistical approach to method extraction in the protein-protein interaction literature","year":2012,"lang":"en","type":"erratum","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Classifier (UML); Computer science; Pattern recognition (psychology); Artificial intelligence; Computational biology; Biology","score_opus":0.05961479667506053,"score_gpt":0.33777227297490175,"score_spread":0.2781574762998412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2610063871","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010271575,0.0038002518,0.012581435,0.10700121,0.85438025,0.00016454823,0.010519514,0.0023140514,0.008211539],"genre_scores_gemma":[0.036188103,0.021222884,0.1036501,0.16764727,0.12485524,0.00096552743,0.05624242,0.009270046,0.47995844],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9929408,0.0009797459,0.0017039651,0.0007461863,0.0033893746,0.00023997504],"domain_scores_gemma":[0.9432934,0.013095987,0.002240241,0.0024797835,0.03781918,0.0010714482],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004658589,0.0013720903,0.0012416762,0.0066393893,0.0030908433,0.0037865508,0.0020809695,0.0029881997,0.035154074],"category_scores_gemma":[0.073891945,0.0009680548,0.0012139776,0.005672788,0.0016901055,0.0030086907,0.0021754433,0.0057196827,0.031031849],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034146702,0.000009278402,0.00016428725,0.00019583525,0.000009973505,0.00022828786,0.000037291567,0.00006804937,0.00012757859,0.000712424,0.9847799,0.013632966],"study_design_scores_gemma":[0.00002308561,0.000020295973,0.0008146476,0.00044254318,0.00003861959,0.00081699254,0.000104944396,0.00045291966,0.00076283823,0.0009664373,0.9955135,0.00004319323],"about_ca_topic_score_codex":0.025584986,"about_ca_topic_score_gemma":0.032427587,"teacher_disagreement_score":0.035154074,"about_ca_system_score_codex":0.0037239934,"about_ca_system_score_gemma":0.007293215,"threshold_uncertainty_score":0.11760211},"labels":[],"label_agreement":null},{"id":"W2610342131","doi":"10.12688/f1000research.11389.2","title":"PubRunner: A light-weight framework for updating text mining results","year":2017,"lang":"en","type":"preprint","venue":"F1000Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Computer science; Upload; Workflow; Word2vec; Biomedical text mining; Data science; Field (mathematics); Domain (mathematical analysis); Information retrieval; World Wide Web; Text mining; Data mining; Artificial intelligence","score_opus":0.07276717202123249,"score_gpt":0.3943506542953737,"score_spread":0.3215834822741412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2610342131","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00156371,0.0006695904,0.69327605,0.0012907635,0.00057351537,0.00082153844,0.015657809,0.2832116,0.0029354058],"genre_scores_gemma":[0.012611302,0.0006694063,0.92359895,0.0005652319,0.00029246707,0.00096057623,0.032232262,0.02570628,0.0033634724],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9794862,0.0050858073,0.0044962675,0.004193583,0.0060058073,0.0007323496],"domain_scores_gemma":[0.9226106,0.037126824,0.0054038437,0.02072153,0.011310136,0.002827004],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04593455,0.0051010256,0.00350367,0.025194963,0.0040961015,0.014307741,0.009368242,0.004213797,0.024581164],"category_scores_gemma":[0.13240795,0.004654617,0.005550288,0.017413579,0.0029281878,0.02846147,0.013520277,0.0057450887,0.02733035],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014493676,0.00047541517,0.0060088974,0.003744162,0.00066129304,0.0010975609,0.002820329,0.009564828,0.007923806,0.061604783,0.32536697,0.57928264],"study_design_scores_gemma":[0.00056904694,0.0002147922,0.002344671,0.0012437593,0.00040132907,0.0008878314,0.00072879513,0.14150994,0.021906178,0.21859641,0.6109844,0.00061285315],"about_ca_topic_score_codex":0.008580081,"about_ca_topic_score_gemma":0.015148645,"teacher_disagreement_score":0.04593455,"about_ca_system_score_codex":0.0027722334,"about_ca_system_score_gemma":0.008165462,"threshold_uncertainty_score":0.24292797},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"not_applicable","genre":"software","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"not_applicable","genre":"software","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"agree"},{"id":"W2611771272","doi":"10.1159/000132722","title":"Listing of forthcoming and published papers available on line","year":2008,"lang":"en","type":"article","venue":"Cytogenetics and Cell Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Listing (finance); Biology; Computational biology; Business","score_opus":0.034643235414971274,"score_gpt":0.2448506783082256,"score_spread":0.21020744289325433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2611771272","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001273887,0.0047965096,0.003119046,0.0040028035,0.013873238,0.0008189552,0.18059976,0.008447056,0.7830689],"genre_scores_gemma":[0.0011025512,0.0018289059,0.0017927807,0.00090794585,0.0016097578,0.00016104788,0.06295156,0.0014940539,0.92815137],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993274,0.000046959754,0.00009089715,0.000118281314,0.0003561344,0.000060377104],"domain_scores_gemma":[0.9954058,0.00046888407,0.00029316478,0.00047376053,0.0016243398,0.0017340434],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009992846,0.0012203697,0.0017372164,0.0066583417,0.00089765573,0.0050759395,0.001124795,0.001035911,0.8559198],"category_scores_gemma":[0.004890833,0.0005623807,0.0007911029,0.007333313,0.00024899477,0.0021725232,0.0016582783,0.0011263244,0.8497693],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006198134,0.000035126162,0.0001099533,0.0004000157,0.000005861701,0.000082384206,0.000008352876,0.00003070644,0.001783943,0.00057306135,0.9576681,0.0392405],"study_design_scores_gemma":[0.000028605942,0.000024568644,0.0005089797,0.0000847807,0.00000397666,0.00006611076,0.000010512374,0.00002059562,0.00031060653,0.00038138073,0.9985514,0.000008470283],"about_ca_topic_score_codex":0.00066527864,"about_ca_topic_score_gemma":0.0017963303,"teacher_disagreement_score":0.8559198,"about_ca_system_score_codex":0.0006960711,"about_ca_system_score_gemma":0.0013070045,"threshold_uncertainty_score":0.20551294},"labels":[],"label_agreement":null},{"id":"W2612992329","doi":"10.2196/medinform.7235","title":"Effective Information Extraction Framework for Heterogeneous Clinical Reports Using Online Machine Learning and Controlled Vocabularies","year":2017,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Information extraction; Controlled vocabulary; Machine learning; Artificial intelligence; Unified Medical Language System; Information retrieval; Data science; Data mining","score_opus":0.02581817604947437,"score_gpt":0.3984720093010376,"score_spread":0.3726538332515632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2612992329","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015234006,0.000867942,0.95299035,0.0010814948,0.000090469235,0.0012959732,0.004780098,0.021876518,0.0017831788],"genre_scores_gemma":[0.082736686,0.0003483817,0.90065926,0.0002621911,0.000082727194,0.00092701864,0.013587569,0.0005291693,0.0008668657],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99164,0.0022154115,0.0023213571,0.0017753894,0.0018135767,0.00023418014],"domain_scores_gemma":[0.97972167,0.01168263,0.0023292252,0.0022131898,0.0036959653,0.00035732193],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009704759,0.001706053,0.0015029773,0.008793053,0.0013394936,0.0036205791,0.0036476138,0.00118135,0.0032907939],"category_scores_gemma":[0.028191943,0.0007574856,0.0024557237,0.0036842402,0.0010217057,0.007264676,0.004126266,0.0016864087,0.0019046266],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005221181,0.00047354444,0.008122111,0.002701074,0.00029812404,0.0015934132,0.002173548,0.03181736,0.041587595,0.023842178,0.030726919,0.8561419],"study_design_scores_gemma":[0.0004147372,0.00042617784,0.0051700454,0.00085267157,0.00044753507,0.0015715415,0.0014327018,0.78007245,0.076436855,0.049573436,0.08327136,0.00033050994],"about_ca_topic_score_codex":0.011668826,"about_ca_topic_score_gemma":0.010079524,"teacher_disagreement_score":0.011668826,"about_ca_system_score_codex":0.0023537928,"about_ca_system_score_gemma":0.005967319,"threshold_uncertainty_score":0.05132425},"labels":[],"label_agreement":null},{"id":"W2616885684","doi":"10.1016/j.jbi.2017.05.016","title":"RysannMD: A biomedical semantic annotator balancing speed and accuracy","year":2017,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Natural language processing; Artificial intelligence; Information retrieval","score_opus":0.017709520062121183,"score_gpt":0.3075889200199805,"score_spread":0.28987939995785933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2616885684","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029577164,0.0028642525,0.7771882,0.0020639922,0.0016393834,0.0009861543,0.015575063,0.16038607,0.009719679],"genre_scores_gemma":[0.06631463,0.0007346321,0.8816465,0.0010425933,0.00017958504,0.000482533,0.029413912,0.009981706,0.010203881],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9841541,0.0054272995,0.001824756,0.0033972648,0.004552616,0.00064384297],"domain_scores_gemma":[0.971806,0.010242001,0.0007163969,0.008883803,0.0074864016,0.0008653994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01947826,0.0024395601,0.0029810357,0.008073088,0.0032251566,0.0052739014,0.0049554477,0.0039943927,0.010286418],"category_scores_gemma":[0.04290352,0.0019623488,0.0029674415,0.004745724,0.001600245,0.010020512,0.011901215,0.0029475624,0.009928339],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030666904,0.0005225586,0.00770314,0.0031996132,0.0011288978,0.00069143943,0.0012566028,0.01195872,0.036559775,0.011548987,0.18109973,0.7412639],"study_design_scores_gemma":[0.00079525134,0.0005082967,0.0059632137,0.00087267254,0.00097430183,0.0024345156,0.0017623848,0.5374481,0.103910975,0.043265846,0.30149823,0.0005661987],"about_ca_topic_score_codex":0.009721412,"about_ca_topic_score_gemma":0.021405341,"teacher_disagreement_score":0.01947826,"about_ca_system_score_codex":0.0021986847,"about_ca_system_score_gemma":0.005575289,"threshold_uncertainty_score":0.103012145},"labels":[],"label_agreement":null},{"id":"W2617001213","doi":"","title":"Section Heading Recognition in Electronic Health Records Using Conditional Random Fields.","year":2014,"lang":"en","type":"article","venue":"Taipei Medical University Repository","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Conditional random field; Named-entity recognition; Computer science; Section (typography); Medical diagnosis; Clinical decision support system; Artificial intelligence; Natural language processing; Health records; Domain (mathematical analysis); Medical record; Information retrieval; Heading (navigation); Machine learning; Masking (illustration); Decision support system; Data science; Health care; Medicine; Engineering","score_opus":0.011494410067773589,"score_gpt":0.2415547741239862,"score_spread":0.2300603640562126,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2617001213","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16229917,0.007913493,0.32304797,0.0041888007,0.00094275264,0.0017508251,0.43711048,0.053110424,0.0096361125],"genre_scores_gemma":[0.3257372,0.0022367367,0.3059714,0.0006666003,0.00040732606,0.0009578021,0.35918254,0.0006019711,0.004238418],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99813193,0.00034238334,0.00038468948,0.00050266867,0.0004916072,0.00014670634],"domain_scores_gemma":[0.9895086,0.0061916807,0.0016480708,0.0010535318,0.001281607,0.00031635773],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028325927,0.0005419767,0.000805228,0.010916792,0.00057699054,0.0014247192,0.0010346235,0.0010613707,0.004957468],"category_scores_gemma":[0.016997926,0.00027694943,0.0009653414,0.008509862,0.000307854,0.0023272266,0.0012124325,0.00088565005,0.0034150728],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001335854,0.00063555985,0.09689065,0.0028723117,0.00036957522,0.0015425389,0.0006734925,0.010895372,0.015997257,0.01157752,0.14157486,0.715635],"study_design_scores_gemma":[0.00046229424,0.0007885361,0.15367396,0.0024609708,0.0012450322,0.003783156,0.0015165404,0.48246104,0.061185382,0.06802577,0.22412124,0.00027608298],"about_ca_topic_score_codex":0.009203803,"about_ca_topic_score_gemma":0.014761836,"teacher_disagreement_score":0.010916792,"about_ca_system_score_codex":0.00096462754,"about_ca_system_score_gemma":0.0037766132,"threshold_uncertainty_score":0.018300414},"labels":[],"label_agreement":null},{"id":"W2619169087","doi":"10.1016/s0992-5945(17)30056-9","title":"La biologie médicale : à l’avant-garde des échanges des données de santé","year":2017,"lang":"fr","type":"article","venue":"Option/Bio","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"World Federation of Science Journalists","funders":"","keywords":"Avant garde; Humanities; Political science; Art; Art history","score_opus":0.07660851905424075,"score_gpt":0.3301115262115185,"score_spread":0.25350300715727775,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2619169087","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017115967,0.053006295,0.5287835,0.3302226,0.003732194,0.00023862257,0.009404467,0.0039687417,0.053527556],"genre_scores_gemma":[0.20580311,0.047338877,0.6970307,0.021902205,0.004006182,0.00030494021,0.009586608,0.000979053,0.013048352],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97482944,0.010238258,0.002994511,0.0018599727,0.009561708,0.00051620574],"domain_scores_gemma":[0.924057,0.059061874,0.002813952,0.005867443,0.0064554,0.0017443689],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030426174,0.0011250451,0.0017225747,0.008444246,0.0020751888,0.014193613,0.002255716,0.0043771197,0.00502195],"category_scores_gemma":[0.061697446,0.0007459821,0.0018649679,0.009930936,0.0066090627,0.025472354,0.0059279413,0.0079470975,0.002125565],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015214292,0.00012637762,0.0056608957,0.0030504127,0.0004059186,0.00066972966,0.0042359736,0.0032025112,0.005814793,0.57755655,0.05021483,0.34890988],"study_design_scores_gemma":[0.000023649061,0.00004645087,0.0036376526,0.0018857627,0.00012468368,0.0010317297,0.0029492076,0.009220388,0.002790269,0.4435863,0.53457904,0.00012480907],"about_ca_topic_score_codex":0.011803818,"about_ca_topic_score_gemma":0.0102607105,"teacher_disagreement_score":0.030426174,"about_ca_system_score_codex":0.0056444774,"about_ca_system_score_gemma":0.013729381,"threshold_uncertainty_score":0.1609109},"labels":[],"label_agreement":null},{"id":"W2619226332","doi":"","title":"Fiches pratiques pour communiquer efficacement en ligne","year":2012,"lang":"fr","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Ligne; Library science; The Internet; Political science; World Wide Web; Computer science; Art","score_opus":0.019025299078187338,"score_gpt":0.30382814103215466,"score_spread":0.2848028419539673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2619226332","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027315326,0.012369639,0.7713164,0.07196423,0.0042878822,0.0009746702,0.0022604424,0.010804282,0.09870721],"genre_scores_gemma":[0.112283394,0.013825175,0.69518083,0.010233665,0.0025100794,0.0011023378,0.0042220047,0.0015400602,0.15910251],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99312544,0.0029514744,0.00064349634,0.0006472899,0.002267131,0.00036512813],"domain_scores_gemma":[0.9806208,0.009608671,0.0010750015,0.0026556656,0.005333434,0.00070644723],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010345211,0.0013517856,0.00081091566,0.00416225,0.0024897936,0.007934803,0.0018739696,0.002991503,0.017262334],"category_scores_gemma":[0.02047256,0.00060797215,0.0014435436,0.002779036,0.002431085,0.006128227,0.0038663468,0.0042061405,0.009366114],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029618756,0.00040543248,0.0043536606,0.0011470089,0.00012626253,0.00079489255,0.007300988,0.003530879,0.0111525,0.13541918,0.07724684,0.75822616],"study_design_scores_gemma":[0.000080264035,0.00029877128,0.0035546469,0.0007508026,0.00014527273,0.0017402421,0.0037563967,0.009074855,0.017417599,0.05846641,0.90456206,0.00015256411],"about_ca_topic_score_codex":0.01143557,"about_ca_topic_score_gemma":0.013509672,"teacher_disagreement_score":0.017262334,"about_ca_system_score_codex":0.0024414908,"about_ca_system_score_gemma":0.0071267975,"threshold_uncertainty_score":0.057748258},"labels":[],"label_agreement":null},{"id":"W26197850","doi":"10.1007/s10654-015-0070-1","title":"Relation extraction from biomedical text","year":2007,"lang":"en","type":"dissertation","venue":"European Journal of Epidemiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Relationship extraction; Discriminative model; Artificial intelligence; Parsing; Generative grammar; Natural language processing; Generative model; Machine learning; Information extraction; Biomedical text mining; Sentence; Boosting (machine learning); Information retrieval; Text mining","score_opus":0.045902507783402956,"score_gpt":0.37262066374104996,"score_spread":0.326718155957647,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W26197850","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057709638,0.018117473,0.7391647,0.0047959755,0.0012165213,0.001877344,0.12986003,0.02185329,0.025405062],"genre_scores_gemma":[0.14715919,0.0074648317,0.6556385,0.0009554373,0.000851488,0.00093422184,0.17744678,0.001209854,0.008339638],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972351,0.00046586688,0.00046463008,0.00095492776,0.0007731695,0.00010631325],"domain_scores_gemma":[0.9933503,0.0035957894,0.0010140581,0.0007644572,0.0011220506,0.00015343694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015057278,0.0021890784,0.0009797427,0.013492475,0.00095955393,0.0019875187,0.0012894849,0.0012489061,0.008376745],"category_scores_gemma":[0.011384108,0.00057218166,0.0016826334,0.008932634,0.0006761527,0.004569238,0.0021221268,0.0012315712,0.008494511],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043133213,0.00018136491,0.0067050876,0.007663978,0.0002867182,0.002918602,0.0011387565,0.007993458,0.053563666,0.02086363,0.065845065,0.8324083],"study_design_scores_gemma":[0.00013960598,0.0003210434,0.026939198,0.0021274108,0.0010118199,0.0058878586,0.002121865,0.1523074,0.09658178,0.12998149,0.5822842,0.00029635732],"about_ca_topic_score_codex":0.0025460422,"about_ca_topic_score_gemma":0.0037656499,"teacher_disagreement_score":0.013492475,"about_ca_system_score_codex":0.0010879573,"about_ca_system_score_gemma":0.0019299082,"threshold_uncertainty_score":0.028023005},"labels":[],"label_agreement":null},{"id":"W2622330800","doi":"10.11575/prism/34791","title":"Creating and Assessing a Subject-Based Blog for Current Awareness within a Cancer Care Environment","year":2013,"lang":"en","type":"article","venue":"PRISM (University of Calgary)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Subject (documents); Computer science; Public relations; Medicine; Political science; World Wide Web","score_opus":0.015362271155558088,"score_gpt":0.24915264201769713,"score_spread":0.23379037086213905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2622330800","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7782684,0.0048549105,0.025663448,0.03365245,0.0045809443,0.0070974603,0.0024788417,0.0022998173,0.14110367],"genre_scores_gemma":[0.8189247,0.0053162635,0.12174204,0.009302204,0.0014211503,0.0040183584,0.0018688742,0.00055019255,0.036856238],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99275416,0.0043550837,0.00048045948,0.00035042706,0.0016667324,0.000393059],"domain_scores_gemma":[0.9493505,0.028352354,0.0038831984,0.002935391,0.0101448065,0.005333739],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.019881787,0.00035323514,0.00032201133,0.0026496865,0.0027042944,0.007834058,0.0009504579,0.0010799175,0.006895024],"category_scores_gemma":[0.04076947,0.00021639133,0.00034930804,0.0019522979,0.0018056735,0.0055103507,0.00446071,0.0013757527,0.0018239269],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024553618,0.0011440845,0.053459127,0.0043069753,0.000046252862,0.00092657236,0.24696483,0.00026830993,0.0054407846,0.0036757991,0.07450336,0.6090184],"study_design_scores_gemma":[0.00010241248,0.0017094625,0.060075045,0.0061598225,0.0001887819,0.00077587843,0.2859917,0.00090505363,0.0036234907,0.0044742795,0.6358173,0.00017680615],"about_ca_topic_score_codex":0.0013157461,"about_ca_topic_score_gemma":0.0055700704,"teacher_disagreement_score":0.9921659,"about_ca_system_score_codex":0.002301955,"about_ca_system_score_gemma":0.0068533327,"threshold_uncertainty_score":0.10514617},"labels":[],"label_agreement":null},{"id":"W2657797703","doi":"10.3389/fmicb.2017.01068","title":"Context Is Everything: Harmonization of Critical Food Microbiology Descriptors and Metadata for Improved Food Safety and Surveillance","year":2017,"lang":"en","type":"article","venue":"Frontiers in Microbiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":56,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Centre for Disease Control; Public Health Agency of Canada; University of Manitoba; University of British Columbia; Simon Fraser University","funders":"Genome British Columbia; Government of Canada; Genome Canada","keywords":"Metadata; Food microbiology; Harmonization; Food safety; Context (archaeology); Clinical microbiology; Business; Biotechnology; Biology; Microbiology; Food science; Computer science; World Wide Web; Bacteria","score_opus":0.017509969376369526,"score_gpt":0.2576362256928161,"score_spread":0.24012625631644655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2657797703","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03349728,0.0049872636,0.85270256,0.01693189,0.001774396,0.004763293,0.025975822,0.01660361,0.04276385],"genre_scores_gemma":[0.1022033,0.0031702346,0.8263146,0.0036353571,0.0004082488,0.0016963453,0.05601798,0.0017780036,0.004775988],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.987765,0.004442249,0.0027584985,0.0017986819,0.0025311883,0.00070435886],"domain_scores_gemma":[0.9660375,0.00597549,0.0035235442,0.01316737,0.009102072,0.0021939895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.030729597,0.00097157987,0.0010230991,0.01006862,0.0018718231,0.007506683,0.0035609924,0.0018794977,0.0025737542],"category_scores_gemma":[0.032273572,0.000727765,0.0015689011,0.010321559,0.0020819828,0.016044106,0.010220032,0.0024489765,0.0018253122],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005753668,0.00061417394,0.021938257,0.0021175456,0.0002490218,0.00061860913,0.005825148,0.0068139546,0.017371798,0.26491082,0.12137159,0.5575937],"study_design_scores_gemma":[0.00011572221,0.00021945708,0.011759649,0.0025124708,0.00027205577,0.00047288986,0.0048328675,0.017613787,0.011273609,0.13012363,0.82054967,0.0002541819],"about_ca_topic_score_codex":0.014733955,"about_ca_topic_score_gemma":0.01175439,"teacher_disagreement_score":0.030729597,"about_ca_system_score_codex":0.0036044507,"about_ca_system_score_gemma":0.01384966,"threshold_uncertainty_score":0.16251558},"labels":[],"label_agreement":null},{"id":"W2684939184","doi":"10.4018/ijsvr.2017010102","title":"Medical Semiotics","year":2017,"lang":"en","type":"article","venue":"International Journal of Semiotics and Visual Rhetoric","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Semiotics; Visual semiotics; Epistemology; Biosemiotics; Semiosis; Social semiotics; Semiotics of culture; Sociology; Philosophy","score_opus":0.021045559061221648,"score_gpt":0.3690658867438439,"score_spread":0.34802032768262225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2684939184","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013928709,0.024696091,0.19490914,0.03841921,0.0030508225,0.00024384989,0.0008348279,0.00046753383,0.72344977],"genre_scores_gemma":[0.7989182,0.013691673,0.092533365,0.008174486,0.0031159145,0.00053687097,0.0011086512,0.0003436954,0.08157714],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9951473,0.0028470794,0.00046656642,0.0005827443,0.00073559396,0.00022065337],"domain_scores_gemma":[0.9953668,0.002872289,0.0003772637,0.0005362704,0.0006035744,0.00024386006],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033252982,0.0006311319,0.00060096435,0.0025862837,0.0031977238,0.006601394,0.0007473137,0.002421895,0.010304069],"category_scores_gemma":[0.006660575,0.0002869736,0.00067282584,0.0015881269,0.02050353,0.0056486297,0.0035059785,0.0024709557,0.0032522317],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000056198214,0.00000290904,0.00005374403,0.000058301383,0.0000022143558,0.000048488208,0.0012242389,0.00008238578,0.00010372788,0.9913373,0.0027320802,0.0043489304],"study_design_scores_gemma":[0.000007157408,0.00001165684,0.00014230072,0.00013545823,0.0000033036215,0.000367426,0.0010377122,0.00051101856,0.00022525586,0.781907,0.21564111,0.000010595207],"about_ca_topic_score_codex":0.0011365825,"about_ca_topic_score_gemma":0.00067696354,"teacher_disagreement_score":0.010304069,"about_ca_system_score_codex":0.002995837,"about_ca_system_score_gemma":0.0018170538,"threshold_uncertainty_score":0.0344705},"labels":[],"label_agreement":null},{"id":"W2733611476","doi":"10.18438/b88s9k","title":"PubMed’s Native Interface Remains the Best Tool for Systematic Searching of its Biomedical Citations","year":2017,"lang":"en","type":"article","venue":"Evidence Based Library and Information Practice","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; MEDLINE; Context (archaeology); Information retrieval; Limiting; Systematic review; World Wide Web; Data science","score_opus":0.03369109689101219,"score_gpt":0.3318680264851003,"score_spread":0.2981769295940881,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2733611476","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003406976,0.38376164,0.049752735,0.15543085,0.032357797,0.026852038,0.26965338,0.017916568,0.06086803],"genre_scores_gemma":[0.019086761,0.42534244,0.2662055,0.044205796,0.016792051,0.07418023,0.11509776,0.009039469,0.03005003],"study_design_codex":"systematic_review","study_design_gemma":"observational","domain_scores_codex":[0.8439871,0.053867485,0.068339236,0.0077698356,0.023406697,0.0026297336],"domain_scores_gemma":[0.3901853,0.36016652,0.07678808,0.02469146,0.13672522,0.011443338],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14133397,0.004126229,0.015353845,0.12113441,0.003347425,0.015723776,0.006512335,0.008140791,0.17897783],"category_scores_gemma":[0.39866522,0.0034036732,0.0052723764,0.10241236,0.0064740623,0.017316274,0.012774363,0.0054678796,0.084165804],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033402134,0.00005050043,0.0006393982,0.46803364,0.0010815095,0.0003210059,0.0012453672,0.00012368648,0.0012851495,0.003879322,0.37332138,0.14968506],"study_design_scores_gemma":[0.00049275375,0.00022196388,0.002275608,0.28902006,0.001294051,0.00047299868,0.00092691224,0.00029612673,0.00065426104,0.0053881085,0.698655,0.00030215512],"about_ca_topic_score_codex":0.003328421,"about_ca_topic_score_gemma":0.007824395,"teacher_disagreement_score":0.85866606,"about_ca_system_score_codex":0.0067517273,"about_ca_system_score_gemma":0.06613278,"threshold_uncertainty_score":0.7474544},"labels":[],"label_agreement":null},{"id":"W2735100420","doi":"10.1097/sla.0000000000002417","title":"DCIS and Breast Cancer: Challenging the Paradigm","year":2017,"lang":"en","type":"letter","venue":"Annals of Surgery","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Women's College Hospital","funders":"","keywords":"Medicine; Breast cancer; Cancer; MEDLINE; Oncology; Gynecology; Internal medicine","score_opus":0.1887701635337099,"score_gpt":0.3490908083008183,"score_spread":0.1603206447671084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2735100420","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003799296,0.0032048218,0.0013974617,0.99161875,0.00254783,0.0000044914673,0.000043165463,0.000010199068,0.0007933905],"genre_scores_gemma":[0.034963325,0.028231664,0.010213738,0.76824576,0.15566167,0.00010232725,0.00019376531,0.00006990728,0.0023178973],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9926819,0.0034679088,0.0010853738,0.00083762757,0.0016141507,0.00031297794],"domain_scores_gemma":[0.8768235,0.10462357,0.003588705,0.002672773,0.009250764,0.0030407635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0146808885,0.00046364963,0.0013663078,0.0017899973,0.0026709563,0.0055301012,0.0021200094,0.014477763,0.00347287],"category_scores_gemma":[0.061445195,0.0005049038,0.0010520993,0.0014235162,0.0069940323,0.016491842,0.0044521103,0.037453663,0.0016990273],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016037466,0.00010824455,0.0046772216,0.0008557951,0.00014578625,0.003439045,0.0010929748,0.00072929286,0.00046573277,0.16773126,0.6259054,0.19468892],"study_design_scores_gemma":[0.00008462732,0.00005029337,0.0015427736,0.0012529909,0.00008190917,0.005645651,0.0022639518,0.0033035614,0.0002335795,0.57694757,0.40850186,0.00009122651],"about_ca_topic_score_codex":0.004175537,"about_ca_topic_score_gemma":0.00756482,"teacher_disagreement_score":0.0146808885,"about_ca_system_score_codex":0.0038375878,"about_ca_system_score_gemma":0.005947006,"threshold_uncertainty_score":0.07764089},"labels":[],"label_agreement":null},{"id":"W2738186538","doi":"10.1111/coin.12125","title":"An open source and modular search engine for biomedical literature retrieval","year":2017,"lang":"en","type":"article","venue":"Computational Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University; Université du Québec à Montréal","funders":"","keywords":"Computer science; Information retrieval; Search engine; Metadata; License; Full text search; World Wide Web; Natural language user interface; Natural language; Table (database); Modular design; Natural language processing; Database; Programming language","score_opus":0.04971705101010006,"score_gpt":0.38567683937158487,"score_spread":0.3359597883614848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2738186538","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013543308,0.007891956,0.36070168,0.0022096827,0.0006036566,0.0022645057,0.19397138,0.38665152,0.032162342],"genre_scores_gemma":[0.053720646,0.0038935866,0.55617857,0.0015955424,0.00028257386,0.0016216277,0.3451876,0.014292798,0.023227062],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982377,0.0003738505,0.00037834354,0.00032389536,0.0005721104,0.00011404254],"domain_scores_gemma":[0.99570835,0.0020785348,0.00036954743,0.00052542967,0.00086999364,0.00044819026],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026968687,0.0013663456,0.0014615024,0.011652589,0.0010064173,0.0030377195,0.0016359478,0.001468584,0.03136673],"category_scores_gemma":[0.011815086,0.0006426981,0.0011551023,0.0072074817,0.00039635258,0.004148209,0.0040715085,0.0008249843,0.026829232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015518484,0.00030733852,0.0027720889,0.0095222,0.0006783843,0.0016069667,0.00063182914,0.0027297663,0.040058047,0.019900747,0.40422955,0.5160112],"study_design_scores_gemma":[0.00068351056,0.00028227124,0.008172067,0.0013465664,0.000475258,0.0027294622,0.00042611017,0.03586391,0.037452284,0.027667137,0.88451886,0.00038254724],"about_ca_topic_score_codex":0.0024263645,"about_ca_topic_score_gemma":0.004521331,"teacher_disagreement_score":0.03136673,"about_ca_system_score_codex":0.0008852556,"about_ca_system_score_gemma":0.002516822,"threshold_uncertainty_score":0.10493219},"labels":[],"label_agreement":null},{"id":"W2739356706","doi":"10.1145/3103010.3103023","title":"Clinically Significant Information Extraction from Radiology Reports","year":2017,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Extraction (chemistry); Radiology; Medicine; Information retrieval; Medical physics; Chromatography","score_opus":0.020720555777530614,"score_gpt":0.32317440344710785,"score_spread":0.3024538476695772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2739356706","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30106968,0.0071796672,0.5719896,0.0029798467,0.0005858049,0.0021723476,0.07222673,0.025567101,0.016229251],"genre_scores_gemma":[0.4253294,0.0025651162,0.50704086,0.00034931203,0.00048102354,0.0005894118,0.059957433,0.00046987543,0.0032176082],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9985764,0.0002531198,0.0003322444,0.0002890187,0.00043865028,0.000110603374],"domain_scores_gemma":[0.99241275,0.0039839004,0.0012763231,0.0005472045,0.0015952305,0.00018460042],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012476774,0.0010481213,0.0007861078,0.007527024,0.0005942907,0.0012182739,0.00088837865,0.0010190629,0.0029227382],"category_scores_gemma":[0.009227069,0.00035737565,0.00091589824,0.0039246106,0.00032522058,0.001203329,0.0009545967,0.000803536,0.0023406602],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000856848,0.0002803783,0.026042663,0.0027143057,0.00012409883,0.0044218535,0.0010277549,0.0052777454,0.12465505,0.004117796,0.03433407,0.79614747],"study_design_scores_gemma":[0.0003242687,0.0009654544,0.16682453,0.0015126995,0.0010950794,0.019165581,0.002942885,0.18686698,0.34852937,0.02703917,0.24430506,0.0004288654],"about_ca_topic_score_codex":0.0024792075,"about_ca_topic_score_gemma":0.0022908887,"teacher_disagreement_score":0.007527024,"about_ca_system_score_codex":0.00062242657,"about_ca_system_score_gemma":0.0020389936,"threshold_uncertainty_score":0.009777546},"labels":[],"label_agreement":null},{"id":"W2740680318","doi":"10.1177/1460458217719560","title":"Context-aware grading of quality evidences for evidence-based decision-making","year":2017,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Grading (engineering); Computer science; Context (archaeology); Quality (philosophy); Process management; Knowledge management; Business; Engineering; Geography","score_opus":0.17928933615977702,"score_gpt":0.47219630009875707,"score_spread":0.2929069639389801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2740680318","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08024375,0.0064310054,0.8997777,0.002954245,0.00024871674,0.0014873787,0.0012413146,0.0019856575,0.0056301137],"genre_scores_gemma":[0.36582968,0.0012561881,0.6311993,0.0001927379,0.00011485244,0.00031663425,0.000638137,0.000037645288,0.0004147792],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9910545,0.0026527415,0.0021883606,0.00089107186,0.0029588416,0.0002544351],"domain_scores_gemma":[0.96927065,0.015567624,0.004729572,0.0025197621,0.0070689847,0.0008434683],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010229321,0.0010137013,0.0013580221,0.011519535,0.0011833614,0.005266822,0.0020387403,0.0014645082,0.0016062411],"category_scores_gemma":[0.052659713,0.0005504133,0.0018588533,0.0049066087,0.0006710783,0.004466482,0.0021273124,0.0017168411,0.00070990407],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005775976,0.0008884705,0.031358846,0.0031745562,0.0007720122,0.00056277093,0.0013117287,0.027809152,0.015912479,0.01453368,0.005760675,0.8973381],"study_design_scores_gemma":[0.00034621675,0.0009389103,0.035321828,0.0032477127,0.0022165363,0.0014463892,0.0018482802,0.73667854,0.057772804,0.13438438,0.025370676,0.00042780038],"about_ca_topic_score_codex":0.0027445403,"about_ca_topic_score_gemma":0.0057445485,"teacher_disagreement_score":0.98977065,"about_ca_system_score_codex":0.001326539,"about_ca_system_score_gemma":0.0031165131,"threshold_uncertainty_score":0.054098487},"labels":[],"label_agreement":null},{"id":"W2741988747","doi":"10.18653/v1/w17-2322","title":"Painless Relation Extraction with Kindred","year":2017,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"Compute Canada","keywords":"Computer science; Relation (database); Relationship extraction; Python (programming language); Task (project management); Biomedical text mining; Data science; Data mining; Information retrieval; Text mining; Programming language","score_opus":0.020222965238287247,"score_gpt":0.3020006754100079,"score_spread":0.28177771017172065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2741988747","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006293595,0.0012595643,0.66946185,0.001423738,0.0007541292,0.00031923768,0.037803132,0.269115,0.013569762],"genre_scores_gemma":[0.04974384,0.0013210526,0.8181841,0.0015248278,0.000277997,0.00068084366,0.07800524,0.034530528,0.015731687],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99634224,0.00061110867,0.00046273245,0.000904956,0.0014501399,0.00022885106],"domain_scores_gemma":[0.99133146,0.0036355818,0.0006228547,0.0031968828,0.0009983934,0.00021486744],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035497337,0.0022759507,0.0013635865,0.005017547,0.0023585137,0.004592876,0.0026750788,0.0014796223,0.027179416],"category_scores_gemma":[0.018514536,0.0018299251,0.0034627109,0.005198837,0.0013601027,0.0063487203,0.005979188,0.0028212578,0.02591759],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048799196,0.00011795714,0.004066833,0.0032226106,0.00029946462,0.0012059071,0.0013160965,0.003803163,0.016896235,0.051840153,0.49485263,0.421891],"study_design_scores_gemma":[0.00007846491,0.000066178574,0.0027107869,0.0004573498,0.00013441475,0.0021327357,0.00032350657,0.045828413,0.038593315,0.1457609,0.7637323,0.0001816638],"about_ca_topic_score_codex":0.002648287,"about_ca_topic_score_gemma":0.006838808,"teacher_disagreement_score":0.027179416,"about_ca_system_score_codex":0.0010204804,"about_ca_system_score_gemma":0.0028025683,"threshold_uncertainty_score":0.0909242},"labels":[],"label_agreement":null},{"id":"W2745161505","doi":"10.12705/664.9","title":"Building the “Plant Glossary”—A controlled botanical vocabulary using terms extracted from the Floras of North America and China","year":2017,"lang":"en","type":"article","venue":"Taxon","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"National Science Foundation","keywords":"Glossary; Categorization; Vocabulary; Taxon; Synonym (taxonomy); Computer science; Taxonomic rank; Natural language processing; Term (time); Set (abstract data type); Linguistics; Artificial intelligence; Interpretation (philosophy); Ecology; Biology","score_opus":0.02429857359007057,"score_gpt":0.2736396031255826,"score_spread":0.24934102953551204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2745161505","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31762534,0.004042365,0.44221887,0.0049257292,0.0010043691,0.008296257,0.16718645,0.0058045555,0.048896044],"genre_scores_gemma":[0.2295395,0.0014157586,0.54753184,0.00069119345,0.00015493638,0.006363654,0.20832676,0.001354417,0.00462195],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99503857,0.0009985295,0.0017888834,0.00089288753,0.0010715805,0.00020949646],"domain_scores_gemma":[0.9844669,0.0060999664,0.0015420151,0.002707098,0.004688329,0.0004957972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007105765,0.0007732032,0.0008971487,0.022502061,0.0024423536,0.003005976,0.0015624951,0.00076632196,0.0039759637],"category_scores_gemma":[0.019492446,0.00068644487,0.0013646653,0.0145010725,0.0013690667,0.0061761322,0.004527033,0.0019572368,0.0011618227],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00031492475,0.0004244676,0.06389033,0.009524072,0.00059412257,0.0018543511,0.041972615,0.007555096,0.09242376,0.09362417,0.13967764,0.54814446],"study_design_scores_gemma":[0.000120337405,0.00016612248,0.09691285,0.005120609,0.0006262528,0.0008118191,0.02106515,0.029473182,0.018119238,0.018840123,0.8083165,0.00042777593],"about_ca_topic_score_codex":0.046009712,"about_ca_topic_score_gemma":0.056642044,"teacher_disagreement_score":0.046009712,"about_ca_system_score_codex":0.0032814108,"about_ca_system_score_gemma":0.008444892,"threshold_uncertainty_score":0.09148383},"labels":[],"label_agreement":null},{"id":"W2746078860","doi":"10.13034/jsst.v10i1.123","title":"Implementation of virtual workflows in KNIME for medicinal chemistry","year":2017,"lang":"en","type":"article","venue":"Journal of Student Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Workflow; Computer science; Data science; Chemist; World Wide Web; Chemistry; Database","score_opus":0.015941998094999948,"score_gpt":0.3743234463669865,"score_spread":0.35838144827198654,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2746078860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0089791315,0.00033641057,0.6657285,0.0007157186,0.00040191482,0.0007702582,0.0101589225,0.30571103,0.0071981302],"genre_scores_gemma":[0.073270604,0.00059052755,0.85492843,0.0008753472,0.00007812605,0.0020074057,0.03679253,0.027037894,0.004419259],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99460524,0.0012440064,0.001111027,0.0012918927,0.0012744811,0.00047329193],"domain_scores_gemma":[0.9882951,0.005261728,0.0007986647,0.0037397463,0.0011612505,0.0007434718],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012461149,0.0016518963,0.00094102684,0.0021688691,0.0011132429,0.004979618,0.003999334,0.0011079467,0.011255073],"category_scores_gemma":[0.017557627,0.001700142,0.0029362065,0.0017495464,0.0015239235,0.004799499,0.0048162392,0.0031874916,0.007716427],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008309017,0.0013808968,0.010113332,0.0057339408,0.0011455785,0.0013984609,0.0039804587,0.053861137,0.05843584,0.12426984,0.16002928,0.5713423],"study_design_scores_gemma":[0.0010448062,0.0005697679,0.0042055305,0.00079211785,0.00021424265,0.00080844195,0.00051486224,0.20757677,0.13151607,0.09782418,0.5542602,0.0006731682],"about_ca_topic_score_codex":0.002377764,"about_ca_topic_score_gemma":0.0020579903,"teacher_disagreement_score":0.012461149,"about_ca_system_score_codex":0.0021700377,"about_ca_system_score_gemma":0.0046931505,"threshold_uncertainty_score":0.0659017},"labels":[],"label_agreement":null},{"id":"W2747006223","doi":"10.3897/tdwgproceedings.1.20486","title":"Darwin Cloud: Mapping real-world data to Darwin Core","year":2017,"lang":"en","type":"article","venue":"Biodiversity Information Science and Standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"National Science Foundation","keywords":"Darwin (ADL); Computer science; Cloud computing; Workflow; Core (optical fiber); Data science; Set (abstract data type); Data mapping; World Wide Web; Database; Software engineering; Programming language","score_opus":0.11242383552824643,"score_gpt":0.35997275007404245,"score_spread":0.24754891454579603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2747006223","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0334842,0.00090409396,0.7144011,0.005745582,0.0010435679,0.0017604993,0.08318918,0.110516414,0.048955366],"genre_scores_gemma":[0.22600298,0.0015080994,0.5533241,0.002151097,0.00026628232,0.0014710675,0.18323758,0.020632293,0.011406528],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99361926,0.001268523,0.0008531239,0.0010850113,0.0027640122,0.0004100573],"domain_scores_gemma":[0.9861223,0.0025290777,0.0009949959,0.006394784,0.0031077478,0.0008510277],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008702655,0.0008355284,0.00072480296,0.004887989,0.001788095,0.005963864,0.0025972938,0.0010091321,0.0062210863],"category_scores_gemma":[0.034828912,0.0007121965,0.0013997363,0.006867322,0.0014846369,0.010253716,0.010982416,0.0021232732,0.0037661078],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011326738,0.00033208812,0.025929444,0.0013397486,0.00037897925,0.0010925307,0.0061190785,0.01685376,0.011578062,0.2736653,0.4260853,0.2354931],"study_design_scores_gemma":[0.00018599024,0.00010284942,0.013162137,0.00047949702,0.00009089067,0.0006894289,0.0030190176,0.08995304,0.0173831,0.19769835,0.67703265,0.00020304167],"about_ca_topic_score_codex":0.021857342,"about_ca_topic_score_gemma":0.016764617,"teacher_disagreement_score":0.021857342,"about_ca_system_score_codex":0.0028345534,"about_ca_system_score_gemma":0.0056276815,"threshold_uncertainty_score":0.04602456},"labels":[],"label_agreement":null},{"id":"W2749169945","doi":"10.3897/tdwgproceedings.1.20637","title":"Using MIxS: An Implementation Report from Two Metagenomic Information Systems","year":2017,"lang":"en","type":"article","venue":"Biodiversity Information Science and Standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"","keywords":"Metadata; Metagenomics; Workflow; Computer science; Ontology; Data science; Human Microbiome Project; Interoperability; Sample (material); Information retrieval; Data mining; World Wide Web; Database; Biology","score_opus":0.05112376328695886,"score_gpt":0.36934142065704473,"score_spread":0.3182176573700859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2749169945","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03594828,0.0018684204,0.59639317,0.0049549397,0.001046154,0.006492971,0.027280366,0.30379224,0.022223456],"genre_scores_gemma":[0.043090828,0.0014853909,0.7852752,0.0018953125,0.00025959083,0.0057351612,0.10150791,0.048045475,0.012705179],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98293024,0.0041575707,0.0022862104,0.0025834003,0.006969011,0.0010734951],"domain_scores_gemma":[0.98602235,0.0044403975,0.0007753058,0.0041288114,0.0028540574,0.0017790366],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028995542,0.0026209247,0.0013373366,0.0027545255,0.0020010313,0.009881124,0.00538549,0.0026100234,0.019374045],"category_scores_gemma":[0.03854559,0.0043098787,0.003650606,0.002867749,0.0014114878,0.011739079,0.01635034,0.0058509866,0.020677019],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006187523,0.0023014518,0.028002195,0.0054808008,0.0016248951,0.0014774909,0.008861581,0.008316079,0.08661434,0.048672635,0.26495764,0.5375034],"study_design_scores_gemma":[0.0017612292,0.0013215493,0.011406932,0.0015797074,0.00063401915,0.0009130673,0.0018936378,0.06838355,0.079163276,0.023746932,0.80821466,0.0009814181],"about_ca_topic_score_codex":0.0073725916,"about_ca_topic_score_gemma":0.003760573,"teacher_disagreement_score":0.9710045,"about_ca_system_score_codex":0.0026964338,"about_ca_system_score_gemma":0.0067507843,"threshold_uncertainty_score":0.15334493},"labels":[],"label_agreement":null},{"id":"W2749305846","doi":"10.1007/978-3-319-60255-4_9","title":"Time Series Analysis for the Most Frequently Mentioned Biomarkers in Breast Cancer Articles","year":2017,"lang":"en","type":"book-chapter","venue":"Studies in big data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Breast cancer; Autoregressive integrated moving average; Series (stratigraphy); Cancer; Oncology; Computer science; Medicine; Time series; Machine learning; Internal medicine; Biology","score_opus":0.13400986005720175,"score_gpt":0.36224689575464203,"score_spread":0.22823703569744028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2749305846","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73794395,0.056487482,0.023714058,0.0043656644,0.0019481941,0.00023086941,0.15655231,0.0018114081,0.016946109],"genre_scores_gemma":[0.8475914,0.016591867,0.033213098,0.00044664642,0.0015592452,0.00039294054,0.08439349,0.0004965903,0.015314604],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9993,0.0001101256,0.00012559969,0.00014876635,0.00025609875,0.000059415837],"domain_scores_gemma":[0.99138707,0.006039122,0.0010987303,0.00016697655,0.0011449773,0.00016314308],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0010325935,0.0003813681,0.00033361735,0.015504368,0.00039728996,0.0014303597,0.00037255778,0.00044220878,0.007822526],"category_scores_gemma":[0.008002324,0.0001062724,0.0010789525,0.015980912,0.00017056789,0.0009318117,0.00041301243,0.0004311701,0.0023802163],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016721252,0.00030168303,0.2399478,0.0050649154,0.0012762586,0.0011889166,0.0011868469,0.003028053,0.021375334,0.0054002837,0.08620961,0.6333483],"study_design_scores_gemma":[0.00006033457,0.0004735979,0.76220036,0.001308191,0.0022814253,0.0031105124,0.003510077,0.03195641,0.020432513,0.008552795,0.16593876,0.00017504283],"about_ca_topic_score_codex":0.0022376827,"about_ca_topic_score_gemma":0.0030423256,"teacher_disagreement_score":0.98449564,"about_ca_system_score_codex":0.00047174018,"about_ca_system_score_gemma":0.00052571885,"threshold_uncertainty_score":0.026168942},"labels":[],"label_agreement":null},{"id":"W2750566992","doi":"10.15265/iy-2017-021","title":"Health Information Management: Changing with Time","year":2017,"lang":"en","type":"review","venue":"Yearbook of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Canadian Institute for Health Information","funders":"Johns Hopkins University","keywords":"Data governance; Information governance; Terminology; Certification; Health informatics; Knowledge management; Data quality; Data management; Analytics; Data science; Information management; Corporate governance; Informatics; Information system; Computer science; Medicine; Management information systems; Business; Engineering; Political science; Data mining; Nursing; Operations management; Public health","score_opus":0.039501979448924286,"score_gpt":0.35414597209191917,"score_spread":0.31464399264299486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2750566992","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012298432,0.1177799,0.021685882,0.67715,0.011952846,0.00007578831,0.00076796283,0.0005685249,0.1577207],"genre_scores_gemma":[0.3431795,0.28969577,0.08055481,0.1569477,0.03147183,0.0002939755,0.0017233501,0.000917926,0.095215224],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9928653,0.0028340798,0.0005909918,0.0009367183,0.002383765,0.00038915806],"domain_scores_gemma":[0.9801937,0.009010657,0.0017314235,0.001402068,0.0051664864,0.002495787],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010695115,0.00042476697,0.00047017008,0.0036923583,0.002679762,0.017692432,0.0015388634,0.0036208855,0.012432248],"category_scores_gemma":[0.017635878,0.00033871166,0.00034770713,0.008426605,0.010214082,0.022397526,0.005075352,0.0054546683,0.003178426],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030808016,0.000064201035,0.0030885565,0.00076966995,0.000027963837,0.0001930254,0.007111197,0.00060178514,0.00058697595,0.32807022,0.21955481,0.4399007],"study_design_scores_gemma":[0.0000032982366,0.000018091132,0.0020654409,0.0006280605,0.000005273845,0.00018078621,0.005121331,0.0002193195,0.0000967642,0.05257196,0.93906486,0.000024780835],"about_ca_topic_score_codex":0.0058481307,"about_ca_topic_score_gemma":0.005920993,"teacher_disagreement_score":0.017692432,"about_ca_system_score_codex":0.0060418067,"about_ca_system_score_gemma":0.008634209,"threshold_uncertainty_score":0.056561828},"labels":[],"label_agreement":null},{"id":"W2751912113","doi":"10.1101/041798","title":"Sharing brain mapping statistical results with the neuroimaging data model","year":2016,"lang":"en","type":"article","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Montreal Neurological Institute and Hospital","funders":"Medical Research Council; National Institutes of Health; Wellcome Trust","keywords":"Neuroimaging; Computer science; Statistical model; Data sharing; Brain mapping; Neuroscience; Data science; Psychology; Artificial intelligence; Medicine; Pathology","score_opus":0.031870919674149296,"score_gpt":0.2505386063030597,"score_spread":0.2186676866289104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2751912113","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032154114,0.00031564766,0.8735796,0.0020022604,0.00028051427,0.000410434,0.053397585,0.061689597,0.00510883],"genre_scores_gemma":[0.071474135,0.0010529392,0.7712603,0.0012262478,0.00031336452,0.002924996,0.12555614,0.022246467,0.003945434],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9875026,0.0051409025,0.0023496107,0.0016973888,0.0030221017,0.00028749704],"domain_scores_gemma":[0.9241844,0.03751313,0.0027284713,0.028063895,0.0068531544,0.00065702933],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.042543303,0.002484611,0.001769122,0.0073277545,0.0011322373,0.009708967,0.0044405963,0.0022781943,0.021312237],"category_scores_gemma":[0.112854734,0.0018138834,0.0060175117,0.005230245,0.0016226912,0.00717498,0.006744785,0.003599864,0.012914939],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013590354,0.00039607933,0.015759697,0.0047993595,0.0020315663,0.0017199243,0.0019489992,0.06279232,0.0077776127,0.2458781,0.38883716,0.2667001],"study_design_scores_gemma":[0.0003460558,0.000134072,0.0037034198,0.0011821978,0.00074739964,0.0008516448,0.0004042898,0.1764517,0.020732695,0.4224251,0.3726773,0.00034414302],"about_ca_topic_score_codex":0.0055906195,"about_ca_topic_score_gemma":0.0044218553,"teacher_disagreement_score":0.9955594,"about_ca_system_score_codex":0.0024193448,"about_ca_system_score_gemma":0.0067895046,"threshold_uncertainty_score":0.22499323},"labels":[],"label_agreement":null},{"id":"W2752320387","doi":"10.1109/tvcg.2017.2745118","title":"PhenoLines: Phenotype Comparison Visualizations for Disease Subtyping via Topic Models","year":2017,"lang":"en","type":"article","venue":"IEEE Transactions on Visualization and Computer Graphics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"SickKids Foundation; Hospital for Sick Children; University of Toronto","funders":"Ontario Genomics; Genome Canada","keywords":"Computer science; Relevance (law); Subtyping; Machine learning; Workflow; Phenotype; Artificial intelligence; Data visualization; Data mining; Data science; Visualization; Biology","score_opus":0.048147966404783756,"score_gpt":0.3388978814107913,"score_spread":0.29074991500600755,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2752320387","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013696167,0.00043577974,0.92306733,0.00082505314,0.00014090525,0.0002752888,0.009947478,0.04870741,0.002904546],"genre_scores_gemma":[0.1693258,0.000623988,0.80973655,0.00027695636,0.00010480626,0.0009893683,0.010790818,0.005822621,0.0023292336],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991704,0.00034401962,0.000091363254,0.00016357326,0.0001798184,0.00005075688],"domain_scores_gemma":[0.99353397,0.0041528423,0.00053112063,0.00081049156,0.0007401205,0.00023153727],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029116785,0.001434883,0.0007043775,0.0044053756,0.00055091875,0.0029349907,0.0012896672,0.0009999151,0.014574338],"category_scores_gemma":[0.014020366,0.0005736,0.0015740546,0.002277491,0.0004467159,0.0030168635,0.0028151271,0.0016247281,0.00232258],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001877755,0.00036033193,0.017919997,0.0019766279,0.00042169966,0.0009318514,0.007165637,0.048345473,0.022477012,0.0683714,0.1342398,0.6959124],"study_design_scores_gemma":[0.00038419376,0.0003025534,0.010626667,0.0005184573,0.00020771357,0.0007285073,0.0015148488,0.6478126,0.022735776,0.14556372,0.16939856,0.0002063347],"about_ca_topic_score_codex":0.004077164,"about_ca_topic_score_gemma":0.005721174,"teacher_disagreement_score":0.014574338,"about_ca_system_score_codex":0.00078158843,"about_ca_system_score_gemma":0.0010338761,"threshold_uncertainty_score":0.048756063},"labels":[],"label_agreement":null},{"id":"W275547703","doi":"","title":"Review of: Incommensurability and Translation: Kuhnian Perspectives on Scientific Communication and Theory Change","year":2001,"lang":"en","type":"article","venue":"Sound Ideas (University of Puget Sound)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Epistemology; Translation (biology); Philosophy; Economics; Positive economics; Neoclassical economics; Sociology; Chemistry","score_opus":0.04126913153002331,"score_gpt":0.2772358963236079,"score_spread":0.23596676479358458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W275547703","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003140759,0.81630874,0.0015135948,0.162938,0.011354728,0.000006668984,0.00005990666,0.000018188686,0.007486098],"genre_scores_gemma":[0.019215187,0.8809974,0.0024508475,0.060235124,0.032085918,0.000048413345,0.00015804126,0.00006504435,0.004743988],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968714,0.0015089805,0.00028642165,0.00039419899,0.0007813522,0.00015771041],"domain_scores_gemma":[0.97285527,0.02197113,0.0010181598,0.000631607,0.0030485238,0.00047534506],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.007996996,0.0008450042,0.0018913499,0.005906828,0.0016920672,0.005901198,0.0027837192,0.00902906,0.007761652],"category_scores_gemma":[0.024078742,0.00044486905,0.0005541526,0.009602979,0.010829798,0.015407239,0.0033786432,0.008213111,0.0027212065],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058300804,0.0000324165,0.00019067286,0.0047430783,0.000060029237,0.00021705692,0.001789902,0.00030798672,0.0001734686,0.20177913,0.524476,0.26617196],"study_design_scores_gemma":[0.000019964516,0.000014840006,0.0005201228,0.003973341,0.000048390844,0.00048227448,0.0007999253,0.0001399923,0.00013163482,0.10721365,0.886618,0.00003789384],"about_ca_topic_score_codex":0.0058701946,"about_ca_topic_score_gemma":0.0087638525,"teacher_disagreement_score":0.99830794,"about_ca_system_score_codex":0.004115087,"about_ca_system_score_gemma":0.0071221567,"threshold_uncertainty_score":0.042292655},"labels":[],"label_agreement":null},{"id":"W2756665674","doi":"10.1093/bioinformatics/btx613","title":"A collaborative filtering-based approach to biomedical knowledge discovery","year":2017,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"Compute Canada","keywords":"Computer science; Knowledge extraction; Strengths and weaknesses; Data science; Singular value decomposition; Graph; Knowledge graph; Data mining; Information retrieval; Machine learning; Artificial intelligence; Theoretical computer science","score_opus":0.021150662301844248,"score_gpt":0.2974989414543733,"score_spread":0.276348279152529,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2756665674","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042462465,0.00075358414,0.9896571,0.0010926193,0.00013692111,0.00035115422,0.00059230556,0.000713991,0.0024559833],"genre_scores_gemma":[0.08615848,0.0009443286,0.9058655,0.0005483309,0.00046653635,0.00056522066,0.0019595441,0.00010429509,0.0033877941],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98372126,0.004803622,0.0013097593,0.004020558,0.005663681,0.0004810957],"domain_scores_gemma":[0.95820373,0.027674498,0.0020866243,0.004291433,0.006835119,0.0009086146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014804626,0.0013515255,0.0029545508,0.014925582,0.0032666652,0.005319368,0.004824879,0.0034306718,0.0046539325],"category_scores_gemma":[0.04358276,0.00091763865,0.0036633967,0.014400766,0.0022399162,0.0037813296,0.0036796168,0.0020875838,0.0025077977],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005121794,0.0007272445,0.010080834,0.001849904,0.0014617101,0.00084274914,0.0018905881,0.098322265,0.00951456,0.06435221,0.019759221,0.79068655],"study_design_scores_gemma":[0.00013488872,0.00025650606,0.0031953398,0.00036485313,0.0006159252,0.0010036892,0.00038428552,0.78396195,0.006957218,0.16799091,0.03496587,0.0001685728],"about_ca_topic_score_codex":0.01609169,"about_ca_topic_score_gemma":0.018303962,"teacher_disagreement_score":0.01609169,"about_ca_system_score_codex":0.002290527,"about_ca_system_score_gemma":0.0053052506,"threshold_uncertainty_score":0.07829529},"labels":[],"label_agreement":null},{"id":"W2757662564","doi":"10.2196/medinform.7641","title":"An Ontology to Improve Transparency in Case Definition and Increase Case Finding of Infectious Intestinal Disease: Database Study in English General Practice","year":2017,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Medical Research Council; Wellcome Trust","keywords":"Medicine; Incidence (geometry); Transparency (behavior); Ontology; Cohort; Disease; Computer science; Internal medicine; Mathematics; Computer security","score_opus":0.03014682675176397,"score_gpt":0.3595435981146865,"score_spread":0.32939677136292256,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2757662564","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9859802,0.000873116,0.0059083304,0.00074189977,0.000020400434,0.0024043045,0.0022512039,0.000030142988,0.0017904373],"genre_scores_gemma":[0.96098644,0.0008888751,0.03039683,0.00070207997,0.000026542348,0.0039394856,0.0026101035,0.000025356443,0.00042432133],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.98401153,0.009898138,0.002920835,0.0012756551,0.0014183426,0.00047548252],"domain_scores_gemma":[0.9172006,0.06265734,0.008515992,0.0046653235,0.005485751,0.0014749331],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023595875,0.00022615862,0.0007133638,0.00474797,0.0011372529,0.0021070864,0.0012731482,0.0008776029,0.0029208532],"category_scores_gemma":[0.07327811,0.00047579897,0.0012048732,0.006422069,0.0009078566,0.003285916,0.0031583149,0.00077999383,0.00025785345],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013122889,0.0022643032,0.802665,0.004273726,0.00045231462,0.0022280007,0.06485456,0.0010504397,0.00077775616,0.0037823652,0.004419541,0.11191967],"study_design_scores_gemma":[0.0014235807,0.0022392543,0.8244718,0.0041857148,0.0016708804,0.004150399,0.11236054,0.01249918,0.0012165887,0.0040057534,0.031489708,0.00028660145],"about_ca_topic_score_codex":0.029791135,"about_ca_topic_score_gemma":0.041137323,"teacher_disagreement_score":0.029791135,"about_ca_system_score_codex":0.004975518,"about_ca_system_score_gemma":0.005539408,"threshold_uncertainty_score":0.1247884},"labels":[],"label_agreement":null},{"id":"W2760618582","doi":"10.1186/s13326-017-0153-x","title":"Semantic annotation in biomedicine: the current landscape","year":2017,"lang":"en","type":"review","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":78,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Biomedicine; Annotation; Current (fluid); Data science; Information retrieval; Semantic annotation; Natural language processing; Artificial intelligence; Bioinformatics; Biology","score_opus":0.08133893116983922,"score_gpt":0.41336825159966073,"score_spread":0.33202932042982153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2760618582","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003454375,0.9832683,0.006288335,0.005293353,0.00072535017,0.000019973331,0.000063697364,0.00008751333,0.0039080535],"genre_scores_gemma":[0.0027673226,0.9877523,0.006076669,0.0016162245,0.0008156365,0.000030048357,0.00012444891,0.000027890814,0.0007894224],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.997474,0.00088862673,0.00032923341,0.00034341132,0.00086327497,0.00010134678],"domain_scores_gemma":[0.9879418,0.00931477,0.0005761648,0.0005042369,0.0013982118,0.00026489366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067767063,0.0011121167,0.001786311,0.0065497034,0.00096856366,0.0044588726,0.002350995,0.0037236786,0.003177388],"category_scores_gemma":[0.0101846205,0.00051354786,0.00082612655,0.009739309,0.0047869775,0.009566657,0.0033495338,0.0033658815,0.0023232894],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040745166,0.000049101527,0.00046071463,0.01925975,0.00007712131,0.00016719145,0.0008611647,0.00057468744,0.0010295392,0.05212199,0.027230024,0.8981279],"study_design_scores_gemma":[0.0000066544494,0.000029187184,0.000816343,0.009287884,0.000071701565,0.00078717264,0.00066653936,0.0005089428,0.00068110245,0.035620667,0.9514818,0.000042067],"about_ca_topic_score_codex":0.0033824167,"about_ca_topic_score_gemma":0.002677995,"teacher_disagreement_score":0.0067767063,"about_ca_system_score_codex":0.0025639504,"about_ca_system_score_gemma":0.006381103,"threshold_uncertainty_score":0.03583908},"labels":[],"label_agreement":null},{"id":"W2763913160","doi":"10.2196/resprot.7904","title":"Knowledge Management Framework for Emerging Infectious Diseases Preparedness and Response: Design and Development of Public Health Document Ontology","year":2017,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Ontology; Computer science; Knowledge management; Upper ontology; Knowledge sharing; Domain knowledge; Data science","score_opus":0.3016131344752973,"score_gpt":0.5681051931511185,"score_spread":0.2664920586758212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2763913160","genre_codex":"methods","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041927355,0.000244816,0.98646796,0.0012682773,0.000071222756,0.001080287,0.0006118824,0.001493947,0.0045689736],"genre_scores_gemma":[0.028345658,0.00036118086,0.96603453,0.0002172622,0.000019806874,0.0007091342,0.0018048964,0.00015575303,0.0023518042],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9953818,0.0011077217,0.0010353642,0.0007436371,0.0014557792,0.000275604],"domain_scores_gemma":[0.9951068,0.0013351279,0.00054887414,0.0007991793,0.0017956106,0.00041436765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008999061,0.0006006798,0.00060125,0.0038319791,0.001899221,0.005262562,0.003135335,0.0016108918,0.0022022799],"category_scores_gemma":[0.009372668,0.0006684529,0.0023362571,0.0027590296,0.0017552251,0.006336303,0.0035941373,0.0027327233,0.00093965774],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013005263,0.0006683185,0.0044227997,0.0014823547,0.00021504576,0.0011884436,0.005212557,0.031080322,0.013062372,0.59701717,0.020457795,0.3250628],"study_design_scores_gemma":[0.00013305094,0.00015685504,0.0023523758,0.0012991399,0.0003954143,0.0014194564,0.0028857253,0.29847077,0.019173516,0.17058139,0.50293356,0.00019886282],"about_ca_topic_score_codex":0.024389805,"about_ca_topic_score_gemma":0.02141068,"teacher_disagreement_score":0.024389805,"about_ca_system_score_codex":0.0039697858,"about_ca_system_score_gemma":0.011961329,"threshold_uncertainty_score":0.04849565},"labels":[],"label_agreement":null},{"id":"W2765375876","doi":"10.12688/f1000research.12234.2","title":"Developing data interoperability using standards: A wheat community use case","year":2017,"lang":"en","type":"preprint","venue":"F1000Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"National Science Foundation of Sri Lanka; Horizon 2020; Biotechnology and Biological Sciences Research Council; European Commission; Agence Nationale de la Recherche; National Science Foundation","keywords":"Interoperability; Metadata; Computer science; Semantic interoperability; Data sharing; Data science; Ontology; Data exchange; Best practice; World Wide Web; Medicine","score_opus":0.5417605375835565,"score_gpt":0.5200049841647696,"score_spread":0.021755553418786877,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2765375876","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32850078,0.002670652,0.4733768,0.07155151,0.00048117008,0.0024225817,0.0006161435,0.0015917822,0.11878862],"genre_scores_gemma":[0.5491925,0.0022921928,0.42561385,0.0056220596,0.00017790812,0.0010505648,0.0018127601,0.0007817424,0.013456447],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"design_other","domain_scores_codex":[0.9055215,0.060009424,0.0066290437,0.0040088794,0.020120548,0.0037106303],"domain_scores_gemma":[0.885159,0.06466247,0.0038715077,0.018248875,0.024472766,0.0035853875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11503063,0.0010130208,0.0009096538,0.0054762354,0.009394334,0.015989417,0.0054030577,0.010764041,0.0025447637],"category_scores_gemma":[0.07486249,0.0013357805,0.0021118643,0.008594692,0.007371901,0.034058988,0.018603228,0.0067604496,0.00080944377],"study_design_candidate":"design_other","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029363154,0.0019997167,0.035354186,0.0014887478,0.00017252451,0.014348539,0.11205558,0.0075035384,0.009099288,0.51919466,0.028952334,0.26953727],"study_design_scores_gemma":[0.00022810744,0.00084500434,0.008200962,0.0025370943,0.0001882214,0.006300985,0.09732672,0.03898334,0.019799333,0.14041045,0.6848069,0.00037292764],"about_ca_topic_score_codex":0.017378792,"about_ca_topic_score_gemma":0.017321907,"teacher_disagreement_score":0.11503063,"about_ca_system_score_codex":0.008208195,"about_ca_system_score_gemma":0.00940439,"threshold_uncertainty_score":0.6083474},"labels":[],"label_agreement":null},{"id":"W2769697870","doi":"10.3390/data2040038","title":"Antibody Exchange: Information Extraction of Biological Antibody Donation and a Web-Portal to Find Donors and Seekers","year":2017,"lang":"en","type":"article","venue":"Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Killam Trusts; National Institute of Mental Health; University of Oxford","keywords":"Computer science; Donation; Web resource; Resource (disambiguation); World Wide Web; Political science; Law","score_opus":0.045460320271140844,"score_gpt":0.36926028536753175,"score_spread":0.3237999650963909,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2769697870","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31388772,0.0053892653,0.3310028,0.003134956,0.00073144777,0.0035155711,0.27297747,0.03657469,0.032786068],"genre_scores_gemma":[0.23292515,0.0013497768,0.4640737,0.0006209008,0.00019972885,0.002003596,0.28674474,0.0006033605,0.011479034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99787045,0.00038497403,0.00039125286,0.0005747724,0.00059913314,0.00017939598],"domain_scores_gemma":[0.9938818,0.002401776,0.0010401491,0.0009016359,0.0013597785,0.00041483142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002399739,0.00075956597,0.0006877283,0.01134992,0.00079300674,0.00144652,0.00084408,0.001416448,0.0031725108],"category_scores_gemma":[0.008795556,0.0003062276,0.00071296457,0.007902154,0.00039138552,0.0026158702,0.0019132538,0.0010169947,0.004215674],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001368327,0.0013960741,0.11972186,0.006283769,0.00024190452,0.005031696,0.002986212,0.005349092,0.038325235,0.013869084,0.18280447,0.6226224],"study_design_scores_gemma":[0.00018939875,0.0005465861,0.11921276,0.0011293204,0.00034068475,0.00533573,0.004765907,0.09166166,0.08106337,0.021574123,0.67398393,0.00019652664],"about_ca_topic_score_codex":0.002283834,"about_ca_topic_score_gemma":0.003347204,"teacher_disagreement_score":0.01134992,"about_ca_system_score_codex":0.0006201099,"about_ca_system_score_gemma":0.002018893,"threshold_uncertainty_score":0.0126912},"labels":[],"label_agreement":null},{"id":"W2770706469","doi":"10.1016/j.media.2017.11.007","title":"The semiotics of medical image Segmentation","year":2017,"lang":"en","type":"article","venue":"Medical Image Analysis","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Robarts Clinical Trials","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Semiotics; Computer science; Sign (mathematics); Segmentation; Interpretation (philosophy); Metaphor; Symbol (formal); Semiosis; Perspective (graphical); Artificial intelligence; Image (mathematics); Image segmentation; Interface (matter); Natural language processing; Human–computer interaction; Linguistics; Mathematics","score_opus":0.012701397737422833,"score_gpt":0.3466525148679864,"score_spread":0.3339511171305636,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2770706469","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018372945,0.0079102395,0.9249979,0.0064927978,0.0005292356,0.00015034828,0.0004210494,0.00038809018,0.040737472],"genre_scores_gemma":[0.43020037,0.005134216,0.5513012,0.00106584,0.0007614826,0.0005594248,0.0007572157,0.00024089521,0.009979385],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9964791,0.0017799116,0.0004095745,0.0005057867,0.0007094781,0.000116071235],"domain_scores_gemma":[0.9933882,0.004490259,0.0004939203,0.0006585098,0.0007355223,0.00023356032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003853895,0.00052336295,0.0006317776,0.004617837,0.0018616661,0.006695305,0.0011605512,0.001994446,0.0025497142],"category_scores_gemma":[0.008837,0.00087490655,0.0016073466,0.0023655747,0.014116453,0.0060658823,0.002548356,0.0022473608,0.00069356913],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001378926,0.000007493838,0.00017277565,0.000088856825,0.000008993232,0.00009955242,0.00071372016,0.00106763,0.00053775276,0.98835003,0.00059349847,0.008345909],"study_design_scores_gemma":[0.000009880639,0.000013247681,0.00025699957,0.00007581498,0.000011251162,0.00024694816,0.00032309233,0.0074199736,0.00052773766,0.97229904,0.018801734,0.000014263633],"about_ca_topic_score_codex":0.0029375972,"about_ca_topic_score_gemma":0.0018238488,"teacher_disagreement_score":0.006695305,"about_ca_system_score_codex":0.0019135316,"about_ca_system_score_gemma":0.0014867745,"threshold_uncertainty_score":0.02038163},"labels":[],"label_agreement":null},{"id":"W2771898558","doi":"10.2196/medinform.8765","title":"Standard Anatomic Terminologies: Comparison for Use in a Health Information Exchange–Based Prior Computed Tomography (CT) Alerting System","year":2017,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine","keywords":"Computed tomography; Health information exchange; Medical physics; Tomography; Medicine; Radiology; Computer science; Health information; Health care","score_opus":0.03697540214023662,"score_gpt":0.3419609349281747,"score_spread":0.30498553278793805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2771898558","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91193944,0.0025698105,0.05257213,0.000932437,0.0006443542,0.004454139,0.005185591,0.001814393,0.019887745],"genre_scores_gemma":[0.8475886,0.0013437227,0.13485157,0.00055738847,0.000153939,0.0035122102,0.010119198,0.00046799463,0.0014052826],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9722722,0.0098613845,0.007830483,0.001368525,0.007992632,0.0006748648],"domain_scores_gemma":[0.79601616,0.119590126,0.026576884,0.015148053,0.039676905,0.0029917147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02650718,0.00057818624,0.00059313327,0.007222287,0.00086592697,0.0034434814,0.0015553879,0.0013709232,0.0023872228],"category_scores_gemma":[0.1828682,0.00033534085,0.001566542,0.004857106,0.0013998229,0.0058769863,0.0032706545,0.0011939132,0.0007962727],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.012519799,0.0022097975,0.4247198,0.00700127,0.0010294802,0.0005759014,0.008979017,0.01725826,0.009207785,0.021260774,0.011566507,0.48367155],"study_design_scores_gemma":[0.0030477038,0.023833076,0.61677414,0.0064436905,0.0044841706,0.0031787618,0.023264559,0.10899426,0.032415073,0.027734023,0.14886534,0.00096527726],"about_ca_topic_score_codex":0.003585573,"about_ca_topic_score_gemma":0.0026707966,"teacher_disagreement_score":0.02650718,"about_ca_system_score_codex":0.004175208,"about_ca_system_score_gemma":0.004194316,"threshold_uncertainty_score":0.14018506},"labels":[],"label_agreement":null},{"id":"W2772527702","doi":"10.1177/0036933017743167","title":"Lingua patientis: new words for patient communication and history taking","year":2017,"lang":"en","type":"article","venue":"Scottish Medical Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University Health Network","funders":"","keywords":"Medicine; Flexibility (engineering); Lexicon; Lingua franca; Exploit; Unified Medical Language System; Vocabulary; Medical history; Linguistics; Artificial intelligence; Computer science; Radiology","score_opus":0.031356418473750654,"score_gpt":0.3132185275549381,"score_spread":0.2818621090811875,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2772527702","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056587376,0.013481749,0.40328032,0.08139139,0.009031512,0.0026268137,0.06366188,0.036772806,0.33316615],"genre_scores_gemma":[0.25108793,0.010535546,0.59198856,0.013980744,0.0022828272,0.0019928156,0.046892934,0.008446148,0.072792515],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974648,0.0008896301,0.0007276123,0.00020464965,0.00058553735,0.0001278483],"domain_scores_gemma":[0.9929426,0.0035181318,0.0008450605,0.00084539707,0.0012427066,0.0006061027],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032597028,0.000768487,0.0005328501,0.0033987905,0.0013190843,0.0037641423,0.0008718561,0.0012555441,0.027591968],"category_scores_gemma":[0.01353141,0.00042821866,0.00068014715,0.0023700013,0.0016051206,0.00790375,0.004275703,0.0022446893,0.01534228],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037335366,0.00013627073,0.006974153,0.0023504794,0.00006308788,0.0019624229,0.012907219,0.0005125859,0.012527157,0.11512282,0.42896017,0.4181103],"study_design_scores_gemma":[0.000020951442,0.000031106116,0.0023396707,0.0004977735,0.000025180387,0.0023212165,0.0020729876,0.0008064905,0.0018102798,0.014421066,0.97558755,0.00006568725],"about_ca_topic_score_codex":0.0037736618,"about_ca_topic_score_gemma":0.0051652556,"teacher_disagreement_score":0.027591968,"about_ca_system_score_codex":0.0019044526,"about_ca_system_score_gemma":0.0040808306,"threshold_uncertainty_score":0.09230435},"labels":[],"label_agreement":null},{"id":"W2773129102","doi":"10.29173/jchla/jabsc.v38i3.29336","title":"Yale MeSH Analyzer","year":2017,"lang":"en","type":"article","venue":"Journal of the Canadian Health Libraries Association / Journal de l Association de bilbiothèques de la santé du Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Health Canada","funders":"","keywords":"Computer science; MEDLINE; Information retrieval; Search engine indexing; Metadata; World Wide Web; Chemistry","score_opus":0.005900698087821761,"score_gpt":0.26076184747420217,"score_spread":0.2548611493863804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2773129102","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005305951,0.002663744,0.062784865,0.0019197942,0.00061057723,0.0013580546,0.5603593,0.26269665,0.10230103],"genre_scores_gemma":[0.024568364,0.0028392933,0.26556966,0.0018231875,0.00043394126,0.0026191154,0.5702395,0.050749592,0.08115741],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99866927,0.00018911166,0.0003405146,0.00025437685,0.00045537003,0.00009137701],"domain_scores_gemma":[0.99233973,0.002971387,0.0006689326,0.000904677,0.0027684344,0.00034672953],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0022643274,0.0016841344,0.0012671738,0.013352289,0.0012671914,0.0040290453,0.0018934762,0.0010031889,0.32909593],"category_scores_gemma":[0.015871936,0.0012229757,0.0013383706,0.009605478,0.0004278402,0.0048754243,0.0028157388,0.001246474,0.13952963],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037179113,0.000044057404,0.0013397352,0.0016254653,0.00007861118,0.00018561493,0.0001161409,0.0001969427,0.002179829,0.0062544043,0.8997042,0.08790325],"study_design_scores_gemma":[0.00018612885,0.000073712865,0.0055167126,0.0006838175,0.00014318504,0.00086648687,0.00021903601,0.0038911984,0.007616381,0.0080461325,0.9726491,0.000108146276],"about_ca_topic_score_codex":0.0062347497,"about_ca_topic_score_gemma":0.011109738,"teacher_disagreement_score":0.9977357,"about_ca_system_score_codex":0.001529763,"about_ca_system_score_gemma":0.0038462528,"threshold_uncertainty_score":0.95696324},"labels":[],"label_agreement":null},{"id":"W2773306025","doi":"10.1186/s13326-017-0163-8","title":"Identifying genotype-phenotype relationships in biomedical text","year":2017,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Automatic summarization; Computer science; Set (abstract data type); Precision and recall; Artificial intelligence; Natural language processing; Recall; Training set; Information retrieval; F1 score; Phenotype; Machine learning; Genotype; Genotype-phenotype distinction; Data mining; Gene; Biology","score_opus":0.05060272485733645,"score_gpt":0.3265227518376841,"score_spread":0.2759200269803476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2773306025","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31613612,0.010913422,0.62420565,0.0037899609,0.0003449175,0.0010193948,0.029424526,0.0070933867,0.007072602],"genre_scores_gemma":[0.5034561,0.0024277251,0.4658162,0.00046276205,0.0004166732,0.0005954167,0.025218457,0.00025201958,0.0013546653],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99606234,0.0013291751,0.0007850851,0.0010371497,0.0007074069,0.00007880763],"domain_scores_gemma":[0.9688028,0.023396319,0.0036574765,0.0013394994,0.0025317161,0.0002722612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003948194,0.00068449054,0.00054817373,0.008369832,0.0007818062,0.0015672736,0.0009787127,0.001151512,0.0025396298],"category_scores_gemma":[0.019652836,0.00020040231,0.0006047923,0.0045962413,0.0006813301,0.002021733,0.0010043622,0.0005911644,0.0013605839],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00074510754,0.00046395432,0.07914022,0.009140457,0.00039850376,0.0032000255,0.002437999,0.015823137,0.08303929,0.0060646734,0.020996371,0.77855027],"study_design_scores_gemma":[0.00019914002,0.0007786267,0.27067056,0.0026210675,0.0014218296,0.016027154,0.0032005496,0.32017314,0.16537811,0.06671237,0.15248464,0.00033283955],"about_ca_topic_score_codex":0.0009834161,"about_ca_topic_score_gemma":0.0014429964,"teacher_disagreement_score":0.008369832,"about_ca_system_score_codex":0.00071234757,"about_ca_system_score_gemma":0.0012862551,"threshold_uncertainty_score":0.020880282},"labels":[],"label_agreement":null},{"id":"W2773351420","doi":"10.2196/medinform.7059","title":"Search and Graph Database Technologies for Biomedical Semantic Indexing: Experimental Analysis","year":2017,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Search engine indexing; Information retrieval; Graph database; Graph; Database; Theoretical computer science","score_opus":0.03028209886436811,"score_gpt":0.3609085883083696,"score_spread":0.33062648944400147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2773351420","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8567081,0.02316878,0.06986393,0.0014396452,0.0009778708,0.0050727176,0.018931145,0.008306417,0.015531326],"genre_scores_gemma":[0.78836995,0.006354039,0.15618744,0.00047687598,0.00031625837,0.0028712412,0.041111365,0.0005684426,0.0037443761],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9866326,0.006603463,0.0014943539,0.0013958677,0.0032115735,0.0006621464],"domain_scores_gemma":[0.9632551,0.027098343,0.0016293713,0.0033507794,0.004146264,0.0005201812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012349615,0.0022861843,0.0014784745,0.007420927,0.0012799784,0.0018137272,0.00189939,0.0016751512,0.0043499507],"category_scores_gemma":[0.031319845,0.00046445028,0.002095149,0.006549304,0.0011713271,0.0047173034,0.0025422801,0.0012372673,0.0017666935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.014325183,0.019662743,0.034370292,0.0107774595,0.0031852503,0.00050875917,0.0010057488,0.051442735,0.026723834,0.004722338,0.041876707,0.79139894],"study_design_scores_gemma":[0.0045050294,0.023563521,0.08996789,0.0011831758,0.0052742944,0.0017904423,0.004157644,0.7171909,0.09918338,0.018229326,0.034402397,0.0005519512],"about_ca_topic_score_codex":0.013378895,"about_ca_topic_score_gemma":0.010768733,"teacher_disagreement_score":0.013378895,"about_ca_system_score_codex":0.002576267,"about_ca_system_score_gemma":0.0019828225,"threshold_uncertainty_score":0.06531179},"labels":[],"label_agreement":null},{"id":"W2773368817","doi":"","title":"Identifying Protein-protein Interactions in Biomedical Literature using Recurrent Neural Networks with Long Short-Term Memory","year":2017,"lang":"en","type":"article","venue":"International Joint Conference on Natural Language Processing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":42,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Recurrent neural network; Benchmark (surveying); Feature engineering; Computer science; Artificial intelligence; Machine learning; Feature (linguistics); Artificial neural network; Long short term memory; Term (time); Deep learning","score_opus":0.04678864868712887,"score_gpt":0.35796626419052296,"score_spread":0.31117761550339407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2773368817","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21906279,0.02058863,0.7394633,0.0018543012,0.00042661373,0.00028650562,0.004654321,0.007838719,0.0058248187],"genre_scores_gemma":[0.79464334,0.0049699764,0.18442899,0.0007007477,0.00045361227,0.00030510864,0.009959215,0.00014144366,0.004397524],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99918276,0.00017851911,0.00010815564,0.00024025212,0.00021692741,0.00007340082],"domain_scores_gemma":[0.99838364,0.0007645624,0.00032015433,0.00014270391,0.000347316,0.000041604606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014215918,0.0010867552,0.00088472734,0.0037098993,0.00039580715,0.00093111023,0.0013888314,0.000978357,0.0012309491],"category_scores_gemma":[0.0045152507,0.00032861152,0.0009205429,0.0034242105,0.00030715333,0.0017572724,0.0009218503,0.0007417423,0.0012526535],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058625685,0.000505438,0.01870609,0.0013940653,0.00092066947,0.0012584737,0.00033396264,0.10040699,0.038725924,0.0034029868,0.013557463,0.8202017],"study_design_scores_gemma":[0.000022355993,0.00015078766,0.005765753,0.00008700196,0.00033495334,0.00043878326,0.000087410306,0.97006637,0.0112596955,0.0067847595,0.0049636704,0.00003845241],"about_ca_topic_score_codex":0.0048051714,"about_ca_topic_score_gemma":0.010735867,"teacher_disagreement_score":0.0048051714,"about_ca_system_score_codex":0.00059894816,"about_ca_system_score_gemma":0.0008767934,"threshold_uncertainty_score":0.009554446},"labels":[],"label_agreement":null},{"id":"W2778788850","doi":"10.3968/10006","title":"Knowledge Mapping Analysis on Text Mining Research of Medicine Related Fields in Different Regions","year":2017,"lang":"en","type":"article","venue":"Cross-cultural communication","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Biomedical text mining; Data science; Field (mathematics); Computer science; Multidisciplinary approach; Knowledge extraction; Information retrieval; Information extraction; Annotation; TRACE (psycholinguistics); Text mining; Data mining; Artificial intelligence","score_opus":0.14179487312639905,"score_gpt":0.4575775997127779,"score_spread":0.31578272658637885,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2778788850","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.838788,0.009054914,0.0767095,0.0024088325,0.00027264937,0.0007140372,0.024111748,0.0015064949,0.046433747],"genre_scores_gemma":[0.91826683,0.0028837996,0.062691234,0.00016722422,0.000119495075,0.0005230779,0.011999709,0.00010283804,0.00324581],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9969829,0.00065394666,0.00065273483,0.0005405891,0.000955696,0.00021421614],"domain_scores_gemma":[0.986974,0.007427992,0.001574046,0.00051135937,0.0032161258,0.00029649382],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.002869042,0.00039469375,0.00050199137,0.04160539,0.001255338,0.00265687,0.0006237,0.00044036438,0.003537585],"category_scores_gemma":[0.016404187,0.00012988856,0.0010102104,0.044745527,0.00042720974,0.0025452662,0.0012579382,0.0003463528,0.00086167414],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003058866,0.00030065386,0.26435342,0.004361535,0.0006048324,0.002110362,0.009207399,0.0057047475,0.0126923015,0.016167227,0.011831665,0.6723599],"study_design_scores_gemma":[0.00006793256,0.00027492386,0.70761013,0.0015711063,0.0020660835,0.0033458555,0.028920332,0.06203919,0.027963774,0.0300442,0.13593438,0.00016199454],"about_ca_topic_score_codex":0.0052777855,"about_ca_topic_score_gemma":0.0040982803,"teacher_disagreement_score":0.9583946,"about_ca_system_score_codex":0.0010080463,"about_ca_system_score_gemma":0.0023306168,"threshold_uncertainty_score":0.015173137},"labels":[],"label_agreement":null},{"id":"W2782223359","doi":"10.1007/978-3-319-73450-7_80","title":"Automatic Extraction and Aggregation of Diseases from Clinical Notes","year":2018,"lang":"en","type":"book-chapter","venue":"Advances in intelligent systems and computing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Extraction (chemistry); Computer science; Information retrieval; Chemistry; Chromatography","score_opus":0.028977729120170024,"score_gpt":0.34216663101336164,"score_spread":0.3131889018931916,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2782223359","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13715188,0.021190595,0.7156403,0.0056077866,0.0015653522,0.0014356694,0.07551797,0.023503125,0.01838735],"genre_scores_gemma":[0.17913327,0.00656457,0.7067106,0.00069404277,0.0007726936,0.0003166696,0.09577501,0.0006633696,0.009369829],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99878734,0.00015657855,0.00024459098,0.0003270646,0.00039193738,0.00009250829],"domain_scores_gemma":[0.99766797,0.0009566724,0.00029444645,0.0003926348,0.0005656621,0.00012259756],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015617164,0.0012033295,0.001091294,0.007866168,0.0006378152,0.0030403172,0.0010554706,0.0008073255,0.003452713],"category_scores_gemma":[0.004668638,0.00047042826,0.0014791039,0.004965488,0.00040650603,0.0016262379,0.0021259785,0.00086658815,0.004206536],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038569598,0.00019994726,0.017170181,0.0010827845,0.0003065278,0.001291254,0.0005439375,0.0023783112,0.03335604,0.004945622,0.047530726,0.89080906],"study_design_scores_gemma":[0.00035174008,0.00055987784,0.13518333,0.0016807031,0.0032436291,0.014295459,0.00227227,0.17718253,0.15441187,0.119431295,0.3910289,0.00035842808],"about_ca_topic_score_codex":0.0027742803,"about_ca_topic_score_gemma":0.005596282,"teacher_disagreement_score":0.007866168,"about_ca_system_score_codex":0.0005600127,"about_ca_system_score_gemma":0.0019794907,"threshold_uncertainty_score":0.011550486},"labels":[],"label_agreement":null},{"id":"W2782565892","doi":"10.1186/s12864-017-4338-6","title":"InfAcrOnt: calculating cross-ontology term similarities using information flow by a random walk","year":2018,"lang":"en","type":"article","venue":"BMC Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":91,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"","keywords":"Term (time); Random walk; Biology; Computational biology; Ontology; Gene ontology; Information flow; Computer science; Information retrieval; Genetics; Mathematics; Statistics; Gene; Physics; Epistemology; Linguistics","score_opus":0.02133868785405676,"score_gpt":0.291117411039584,"score_spread":0.26977872318552726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2782565892","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0877817,0.0022205613,0.89435613,0.00037918036,0.00019108874,0.0005182465,0.002347262,0.010530248,0.0016755619],"genre_scores_gemma":[0.3783826,0.0005444757,0.608941,0.00038452604,0.0001866646,0.0007093214,0.00787678,0.00089154835,0.0020830089],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973974,0.0006222085,0.00022347318,0.00092516263,0.0006556614,0.00017601658],"domain_scores_gemma":[0.9935448,0.004289717,0.0006644278,0.00056176935,0.00075188646,0.00018739294],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030530423,0.0021005536,0.0021282893,0.010098892,0.0014394753,0.001736697,0.00251366,0.0023849064,0.0024606627],"category_scores_gemma":[0.013158087,0.00071804924,0.0029356454,0.004460422,0.0009763076,0.0025469651,0.0015741535,0.0017261484,0.001005859],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011031324,0.00071954675,0.039777014,0.001239873,0.0012034212,0.0007189548,0.0005433491,0.4152321,0.0133422725,0.012283906,0.016757797,0.4970786],"study_design_scores_gemma":[0.000035809815,0.00008404131,0.0011362908,0.000028563338,0.000048309677,0.00015005964,0.000033234974,0.98847884,0.0016150944,0.006979447,0.0013869166,0.000023370349],"about_ca_topic_score_codex":0.012991849,"about_ca_topic_score_gemma":0.015719157,"teacher_disagreement_score":0.012991849,"about_ca_system_score_codex":0.0016551876,"about_ca_system_score_gemma":0.002169638,"threshold_uncertainty_score":0.025832474},"labels":[],"label_agreement":null},{"id":"W2782833738","doi":"","title":"Research Guides: Finding Theses and Dissertations: Need Help?","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Data science; Political science; Library science","score_opus":0.1287584397362548,"score_gpt":0.3970640753438233,"score_spread":0.2683056356075685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2782833738","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008038948,0.011818545,0.2401538,0.18537626,0.008842998,0.0016792268,0.046311326,0.086053565,0.41172528],"genre_scores_gemma":[0.022906886,0.014731955,0.47068387,0.012960239,0.0018132024,0.0011696585,0.026866902,0.030897304,0.41797003],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9925282,0.0025108682,0.001372803,0.0006790438,0.0026470001,0.000262016],"domain_scores_gemma":[0.92012334,0.033615295,0.00448834,0.013114675,0.021841403,0.006817024],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01508872,0.0010003819,0.001120863,0.010546498,0.0039026022,0.013028564,0.0027876757,0.0023675212,0.13915068],"category_scores_gemma":[0.113159426,0.0017799184,0.0006272012,0.019971583,0.0027958187,0.0149881,0.004789654,0.0038276752,0.16260599],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048598933,0.000026171294,0.0004756355,0.00037127605,0.0000095014675,0.00012921063,0.0014868515,0.00005447081,0.0004203271,0.009473229,0.76514643,0.22235824],"study_design_scores_gemma":[0.000021289996,0.0000112605,0.0005688071,0.00036931428,0.000008281866,0.00029650118,0.0018391558,0.0001737754,0.0005281583,0.009712689,0.9864395,0.00003120662],"about_ca_topic_score_codex":0.011734311,"about_ca_topic_score_gemma":0.034280267,"teacher_disagreement_score":0.13915068,"about_ca_system_score_codex":0.0029145859,"about_ca_system_score_gemma":0.01259304,"threshold_uncertainty_score":0.46550542},"labels":[],"label_agreement":null},{"id":"W2783145722","doi":"10.1007/s10545-017-0125-4","title":"Text‐based phenotypic profiles incorporating biochemical phenotypes of inborn errors of metabolism improve phenomics‐based diagnosis","year":2018,"lang":"en","type":"article","venue":"Journal of Inherited Metabolic Disease","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; BC Cancer Agency; BC Children's Hospital; University of British Columbia","funders":"European Commission; Genome British Columbia; Michael Smith Health Research BC; Compute Canada; Canadian Institutes of Health Research; Genome Canada","keywords":"Phenomics; Phenotype; Computational biology; Bioinformatics; Medicine; Biology; Genetics; Genomics; Gene; Genome","score_opus":0.01430182054094183,"score_gpt":0.26803577518264854,"score_spread":0.2537339546417067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2783145722","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2816248,0.0034730667,0.3677355,0.005978897,0.00036314805,0.0009360063,0.29687977,0.022822693,0.02018618],"genre_scores_gemma":[0.48597905,0.002571177,0.35160497,0.0011848117,0.00018803489,0.00039214775,0.15442957,0.0008353628,0.002814815],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99896276,0.00023043342,0.00024754406,0.00021502715,0.00029402444,0.000050135848],"domain_scores_gemma":[0.99330586,0.003553357,0.0010886262,0.00076719356,0.00107029,0.00021463486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013699562,0.0009972217,0.0005471501,0.0066683693,0.00038300196,0.0018176283,0.00069620577,0.0008558757,0.0044645527],"category_scores_gemma":[0.011347792,0.00019664259,0.0008715608,0.003821538,0.00027460675,0.0026740462,0.0013139115,0.00083430053,0.0019536195],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016328411,0.0009693333,0.20018171,0.0040973015,0.0007910932,0.004356606,0.0011517272,0.034435023,0.058443457,0.018765623,0.06534196,0.6098335],"study_design_scores_gemma":[0.00025879394,0.00043795374,0.23772964,0.001751748,0.0013612527,0.006524386,0.0016921206,0.34535396,0.079843536,0.09664504,0.22803193,0.00036961233],"about_ca_topic_score_codex":0.003936553,"about_ca_topic_score_gemma":0.005340297,"teacher_disagreement_score":0.0066683693,"about_ca_system_score_codex":0.0006434206,"about_ca_system_score_gemma":0.0012604473,"threshold_uncertainty_score":0.014935434},"labels":[],"label_agreement":null},{"id":"W2785876892","doi":"10.5206/wurjhns.2017-18.16","title":"Western Faculty Profile: Dr. Shannon L. Sibbald","year":2017,"lang":"en","type":"article","venue":"Western Undergraduate Research Journal Health and Natural Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Pulmonary disease; Health care; Medical education; COPD; Family medicine; Medicine; Knowledge translation; Political science; Knowledge management; Internal medicine; Computer science","score_opus":0.1377943538896322,"score_gpt":0.48027486665796937,"score_spread":0.34248051276833713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2785876892","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004929088,0.0054817414,0.00074659043,0.8529866,0.05157216,0.00026817853,0.0013105682,0.00096679205,0.08173824],"genre_scores_gemma":[0.01894114,0.004636513,0.00045753192,0.323196,0.01494049,0.00024025222,0.00050320785,0.00043192433,0.63665307],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992336,0.000100595316,0.00004173994,0.00013160141,0.00036032608,0.00013213335],"domain_scores_gemma":[0.9906002,0.00035957224,0.0002362871,0.000119921606,0.0031024583,0.005581571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001411703,0.00047905216,0.0006902892,0.00090701674,0.0038656471,0.0028499258,0.00091768196,0.00296633,0.20320801],"category_scores_gemma":[0.007870605,0.0003675298,0.0002446775,0.0008911216,0.0007254197,0.0025200215,0.00259208,0.0040444843,0.08505424],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000005109591,0.000011220327,0.00017021652,0.000012837415,5.0314173e-7,0.000055631084,0.00010090151,0.0000039405613,0.000038214675,0.00009265788,0.9957327,0.0037760888],"study_design_scores_gemma":[0.000011326247,0.00003697033,0.0014098439,0.000098042794,0.00000295971,0.00042523872,0.001969991,0.000052896936,0.000105366096,0.00022854407,0.9956398,0.000019056266],"about_ca_topic_score_codex":0.009201281,"about_ca_topic_score_gemma":0.022091728,"teacher_disagreement_score":0.20320801,"about_ca_system_score_codex":0.0022289369,"about_ca_system_score_gemma":0.006347015,"threshold_uncertainty_score":0.67979854},"labels":[],"label_agreement":null},{"id":"W2785982856","doi":"","title":"Ambiguities in medical bitemporalized relational databases: A referent tracking view","year":2017,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Referent; Relational database; Computer science; Relational model; Tracking (education); Relational database management system; Information retrieval; Basis (linear algebra); Database; Artificial intelligence; Mathematics; Psychology; Linguistics","score_opus":0.059306316965160914,"score_gpt":0.3130547663944934,"score_spread":0.2537484494293325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2785982856","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020459594,0.0050845086,0.946881,0.010200014,0.00033300035,0.00024297477,0.0032985106,0.003307914,0.010192475],"genre_scores_gemma":[0.32544863,0.005049269,0.65173566,0.0024490247,0.00076801603,0.00024310916,0.007060126,0.0015948499,0.0056513497],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9765885,0.005368289,0.003367023,0.0034692753,0.010219578,0.0009873719],"domain_scores_gemma":[0.94038624,0.032812215,0.0050456137,0.013201021,0.0076070554,0.0009478109],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020379314,0.0008641972,0.0023918322,0.014657455,0.0035736219,0.019461637,0.0060629128,0.004697126,0.007357693],"category_scores_gemma":[0.08016116,0.0027310045,0.0027295828,0.020831564,0.0069578756,0.038076192,0.010662622,0.005431923,0.001545512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025971848,0.0001005106,0.0036535023,0.0006568519,0.00011639293,0.00094121235,0.003481108,0.01050741,0.0019965821,0.81799704,0.010292094,0.14999758],"study_design_scores_gemma":[0.00004187816,0.00004519709,0.0009622768,0.00047554064,0.00019637607,0.0009739393,0.0018228389,0.09049137,0.007539522,0.837563,0.05978283,0.000105248255],"about_ca_topic_score_codex":0.011015757,"about_ca_topic_score_gemma":0.004725778,"teacher_disagreement_score":0.020379314,"about_ca_system_score_codex":0.003748904,"about_ca_system_score_gemma":0.004132592,"threshold_uncertainty_score":0.10777742},"labels":[],"label_agreement":null},{"id":"W2786310226","doi":"","title":"A taxonomy of disposition-parthood","year":2017,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Mereology; Disposition; Axiom; Taxonomy (biology); Computer science; Decomposition; Epistemology; Mathematical economics; Algebra over a field; Mathematics; Philosophy; Pure mathematics; Ecology","score_opus":0.02148211520578329,"score_gpt":0.2517753753910883,"score_spread":0.230293260185305,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786310226","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036905643,0.002865644,0.779356,0.0048044636,0.00069087115,0.00045234174,0.0035591146,0.0013688966,0.169997],"genre_scores_gemma":[0.48819426,0.0032026723,0.46718237,0.0008927772,0.00050190475,0.00081968436,0.0068495586,0.00072138,0.03163538],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99345654,0.0020817444,0.0010826327,0.0015062822,0.0012249504,0.00064786617],"domain_scores_gemma":[0.9888736,0.0049698125,0.0007568085,0.0026289104,0.0020792498,0.000691652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004711596,0.0008965258,0.0010972692,0.005095736,0.0061583854,0.01010793,0.0025972654,0.0026792486,0.024416478],"category_scores_gemma":[0.015397867,0.0012169611,0.0024664856,0.011825718,0.008582916,0.025278458,0.005072848,0.004634456,0.0056851627],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000010911241,0.000006992591,0.0005728133,0.00003968704,0.0000055462265,0.00009666414,0.0008926482,0.00016778281,0.0001279775,0.98389524,0.0019292069,0.01225448],"study_design_scores_gemma":[0.000008399883,0.000012967173,0.00039157397,0.000059988255,0.000014580288,0.0004631704,0.0007773395,0.0011568267,0.00018521649,0.9461378,0.050777543,0.000014622258],"about_ca_topic_score_codex":0.0035100288,"about_ca_topic_score_gemma":0.002203436,"teacher_disagreement_score":0.024416478,"about_ca_system_score_codex":0.0026189387,"about_ca_system_score_gemma":0.0027901293,"threshold_uncertainty_score":0.08168125},"labels":[],"label_agreement":null},{"id":"W2786476437","doi":"","title":"PubMedReco: A Real-Time Recommender System for PubMed Citations.","year":2017,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Recommender system; World Wide Web; Interface (matter); Information retrieval; Window (computing); Conversation; User interface; Multimedia","score_opus":0.04300958793450745,"score_gpt":0.2695530656552363,"score_spread":0.22654347772072883,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2786476437","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047828928,0.0152557595,0.38692653,0.008644611,0.0031151366,0.0036494385,0.11476447,0.3947084,0.025106788],"genre_scores_gemma":[0.10558714,0.0043258234,0.74684113,0.0027116472,0.0010381025,0.0017726904,0.11027823,0.0064069955,0.021038244],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9968321,0.0009390801,0.00053450273,0.0005294376,0.001047107,0.000117773714],"domain_scores_gemma":[0.9781691,0.012149646,0.001654411,0.0023889723,0.0041944613,0.0014433725],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0069356556,0.0016283976,0.0014517737,0.009909117,0.0011861749,0.0026750294,0.002456525,0.0025112557,0.013129332],"category_scores_gemma":[0.029605951,0.0009963438,0.0010918493,0.0057566892,0.00022353271,0.0037226,0.0020258173,0.0013426661,0.011692895],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016466979,0.00057169405,0.011831697,0.0046275393,0.0008809714,0.0009749897,0.000922798,0.0040675825,0.020871477,0.0038785392,0.5290209,0.42070514],"study_design_scores_gemma":[0.0010983248,0.00083085266,0.015916424,0.00082831975,0.0008867708,0.0017857656,0.0007378886,0.18287927,0.022689123,0.012011173,0.7596315,0.00070450694],"about_ca_topic_score_codex":0.0079421485,"about_ca_topic_score_gemma":0.025422068,"teacher_disagreement_score":0.99732494,"about_ca_system_score_codex":0.00088730984,"about_ca_system_score_gemma":0.0023127934,"threshold_uncertainty_score":0.043921947},"labels":[],"label_agreement":null},{"id":"W2789509918","doi":"10.1017/s1551929500055206","title":"Fostering LIMS Development Through Open Standards Part II - Ontologies and Business Process","year":2006,"lang":"en","type":"article","venue":"Microscopy Today","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nutrasource","funders":"","keywords":"Focus (optics); Process (computing); Work (physics); Engineering; Computer science; Engineering management; Software engineering; Mechanical engineering","score_opus":0.03050492147145667,"score_gpt":0.3400271157208271,"score_spread":0.30952219424937044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2789509918","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015217596,0.0008101647,0.91608113,0.030976921,0.00053494674,0.00039723283,0.000113284135,0.004988,0.03088079],"genre_scores_gemma":[0.17021427,0.0016274696,0.80638206,0.0033829957,0.00049777416,0.00049086846,0.00072294154,0.0011070841,0.015574481],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9749487,0.0118530495,0.0031198338,0.0017424067,0.0070596454,0.0012763711],"domain_scores_gemma":[0.9582704,0.01550495,0.0031900923,0.013099749,0.007535126,0.002399581],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.052490935,0.00043253272,0.0004758768,0.002351741,0.0020791416,0.013189558,0.0030022087,0.0026335584,0.0039055638],"category_scores_gemma":[0.034732193,0.00074788684,0.0009851477,0.002589926,0.0048731025,0.025029307,0.012167986,0.0051439945,0.0018050977],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006214271,0.00037334336,0.0030219064,0.0003861013,0.000043862892,0.00031236277,0.0047412827,0.004462001,0.007169198,0.68499845,0.025494138,0.26893526],"study_design_scores_gemma":[0.000046033176,0.00013389936,0.0017795487,0.0007209516,0.00004574783,0.000513995,0.003302678,0.03855191,0.01668707,0.3246286,0.61348873,0.000100835976],"about_ca_topic_score_codex":0.0015178004,"about_ca_topic_score_gemma":0.0015445704,"teacher_disagreement_score":0.052490935,"about_ca_system_score_codex":0.0041812654,"about_ca_system_score_gemma":0.0073577655,"threshold_uncertainty_score":0.2776019},"labels":[],"label_agreement":null},{"id":"W2790121460","doi":"10.2196/medinform.8175","title":"Representation of Time-Relevant Common Data Elements in the Cancer Data Standards Repository: Statistical Evaluation of an Ontological Approach","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Institute of Allergy and Infectious Diseases; National Cancer Institute","keywords":"Computer science; Data science; Data sharing; Data quality; Dimension (graph theory); Representation (politics); Domain (mathematical analysis); Data integration; Data mining; Service (business); Medicine; Pathology","score_opus":0.10580009607538754,"score_gpt":0.4493720295632368,"score_spread":0.34357193348784926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2790121460","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7144868,0.00036215447,0.27270138,0.00084612204,0.00007864189,0.0019990117,0.005253619,0.0013347713,0.002937589],"genre_scores_gemma":[0.7645732,0.00013687642,0.2248898,0.00010750009,0.000015908125,0.0015708742,0.008320137,0.00012929768,0.00025634875],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.97246474,0.013502441,0.003984206,0.0028865898,0.00643305,0.0007291064],"domain_scores_gemma":[0.8072685,0.14671184,0.013088197,0.013256888,0.018529247,0.0011453547],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.044226952,0.00060556066,0.0005598056,0.008527382,0.0014345831,0.0031308893,0.0017715422,0.0010849977,0.00089338643],"category_scores_gemma":[0.1735674,0.00042547393,0.002000741,0.009665099,0.0013140155,0.003964981,0.00433308,0.0015500919,0.00017102552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027597318,0.001983843,0.52198815,0.0017538385,0.0013032632,0.0007758453,0.008478712,0.06077001,0.008951177,0.024184128,0.006186342,0.36086503],"study_design_scores_gemma":[0.0003302981,0.00076621,0.19238025,0.00055954704,0.0011081769,0.00054670195,0.009241244,0.74255985,0.015800074,0.022478925,0.01399648,0.00023222537],"about_ca_topic_score_codex":0.020208191,"about_ca_topic_score_gemma":0.017299196,"teacher_disagreement_score":0.95577306,"about_ca_system_score_codex":0.003573282,"about_ca_system_score_gemma":0.003962558,"threshold_uncertainty_score":0.23389733},"labels":[],"label_agreement":null},{"id":"W2792141165","doi":"10.6028/nist.ir.7774","title":"NIST workshop on ontology evaluation","year":2011,"lang":"en","type":"report","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Pacific Northwest National Laboratory; National Institute of Diabetes and Digestive and Kidney Diseases; National Institute of Mental Health; National Institute of Standards and Technology; National Institutes of Health; Rensselaer Polytechnic Institute; University of Washington; University of Toronto; Boeing","keywords":"Ontology; Computer science; Brainstorming; Ontology engineering; Process ontology; NIST; Upper ontology; Data science; Quality (philosophy); Knowledge management; Knowledge base; Suggested Upper Merged Ontology; World Wide Web; Artificial intelligence; Domain knowledge; Natural language processing","score_opus":0.1533625971217804,"score_gpt":0.38975211717892855,"score_spread":0.23638952005714814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2792141165","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017526861,0.031910744,0.43458554,0.13235241,0.016682617,0.002782984,0.008283321,0.008479228,0.34739637],"genre_scores_gemma":[0.1276041,0.02005529,0.5874829,0.012340076,0.0042681224,0.0028121027,0.027159091,0.0053456267,0.21293266],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9189735,0.02323338,0.00617451,0.0057999673,0.043179087,0.0026395752],"domain_scores_gemma":[0.9121941,0.019921487,0.0018764463,0.0124907065,0.049105335,0.004411944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08121018,0.0016060002,0.0020389375,0.009901331,0.0059848386,0.015746929,0.0046239165,0.00556569,0.01091605],"category_scores_gemma":[0.09316009,0.0014476465,0.0023776926,0.0059621707,0.004356093,0.016864637,0.009298238,0.007181157,0.007821345],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016606279,0.0004181341,0.00125876,0.00042808594,0.00007900762,0.00012959949,0.0009223331,0.0022984145,0.00239971,0.13621187,0.5050613,0.35062665],"study_design_scores_gemma":[0.000040841267,0.00008457448,0.001993965,0.0007594106,0.00004897541,0.00017516203,0.00075274345,0.0058450773,0.007935247,0.0814598,0.9008259,0.000078331555],"about_ca_topic_score_codex":0.036849666,"about_ca_topic_score_gemma":0.029226612,"teacher_disagreement_score":0.08121018,"about_ca_system_score_codex":0.014782261,"about_ca_system_score_gemma":0.034737226,"threshold_uncertainty_score":0.42948562},"labels":[],"label_agreement":null},{"id":"W2793303269","doi":"10.1109/tcbb.2018.2817488","title":"Automated ICD-9 Coding via A Deep Learning Approach","year":2018,"lang":"en","type":"article","venue":"IEEE/ACM Transactions on Computational Biology and Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":145,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Central South University; National Natural Science Foundation of China","keywords":"Deep learning; Artificial intelligence; Computer science; Convolutional neural network; ENCODE; Support vector machine; Machine learning; Coding (social sciences); Artificial neural network; Pattern recognition (psychology); Data mining","score_opus":0.016495824943256028,"score_gpt":0.28446123188998507,"score_spread":0.26796540694672905,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2793303269","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07941265,0.0008740274,0.8895636,0.0027551504,0.00037583223,0.0005734385,0.0065192925,0.01260159,0.007324457],"genre_scores_gemma":[0.41745415,0.0005574201,0.5559361,0.00082909816,0.00021507312,0.00033579933,0.01732018,0.00025048573,0.0071016997],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982992,0.0003851005,0.00020564956,0.00040904444,0.00050182577,0.00019914434],"domain_scores_gemma":[0.9966587,0.0012132339,0.00038953443,0.0005362963,0.0010663313,0.0001358967],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016753427,0.0009393159,0.00081082876,0.0043345448,0.00068824657,0.0012879652,0.0015191472,0.0010772203,0.0036855603],"category_scores_gemma":[0.0073552486,0.0003445264,0.00067412236,0.0027176787,0.00044844853,0.0016770845,0.0017541752,0.0024235086,0.0022981577],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030973856,0.00044990148,0.008259813,0.00018098154,0.00005760582,0.00018296861,0.00020276452,0.043578617,0.0060370783,0.0069670915,0.029633725,0.90413976],"study_design_scores_gemma":[0.000054884305,0.000073769734,0.0029344945,0.000065594264,0.000030264613,0.0001624127,0.00017540313,0.95235085,0.0080835745,0.027193487,0.008836229,0.00003904978],"about_ca_topic_score_codex":0.010645106,"about_ca_topic_score_gemma":0.015231361,"teacher_disagreement_score":0.010645106,"about_ca_system_score_codex":0.0017007184,"about_ca_system_score_gemma":0.0031844266,"threshold_uncertainty_score":0.021166265},"labels":[],"label_agreement":null},{"id":"W2794604371","doi":"10.1101/262790","title":"Transfer learning for biomedical named entity recognition with neural networks","year":2018,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto; University of New Brunswick","funders":"National Institutes of Health; Nvidia","keywords":"Computer science; Transfer of learning; Named-entity recognition; Artificial intelligence; Deep learning; Annotation; Conditional random field; Task (project management); Artificial neural network; Field (mathematics); Natural language processing; Machine learning; Word (group theory); Noise (video); State (computer science); Deep neural networks; Domain (mathematical analysis); Mathematics","score_opus":0.016855415850946784,"score_gpt":0.2328407093391094,"score_spread":0.2159852934881626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2794604371","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16195157,0.0052109836,0.79140013,0.0020356337,0.0007126911,0.0004405806,0.0026798912,0.028241972,0.007326534],"genre_scores_gemma":[0.7886421,0.0010622194,0.19331008,0.0006208855,0.00029792712,0.00047538156,0.008467924,0.00031020227,0.006813291],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9981833,0.00058209046,0.00013518843,0.0005775666,0.00036452993,0.0001573608],"domain_scores_gemma":[0.9952873,0.0025097858,0.00036708108,0.00083327055,0.0008993072,0.00010322813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004151504,0.0017246922,0.0010900602,0.0025353548,0.00059234066,0.00114057,0.0026207692,0.0020669901,0.0040826625],"category_scores_gemma":[0.008821734,0.00046406814,0.0010260557,0.002472862,0.0009047786,0.0035969352,0.0021909238,0.0022731642,0.0023834184],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033546344,0.0003359738,0.002544477,0.00030080375,0.00021795469,0.0002578521,0.000104387495,0.37840727,0.0062071104,0.0035948022,0.011617062,0.5960768],"study_design_scores_gemma":[0.000012526194,0.000047533165,0.00053887494,0.000020297475,0.000020176682,0.00002955105,0.000033428896,0.9840447,0.0067263013,0.0068949945,0.001616587,0.000015155698],"about_ca_topic_score_codex":0.0064648096,"about_ca_topic_score_gemma":0.005024794,"teacher_disagreement_score":0.0064648096,"about_ca_system_score_codex":0.0018465898,"about_ca_system_score_gemma":0.0011366258,"threshold_uncertainty_score":0.02195555},"labels":[],"label_agreement":null},{"id":"W2794792323","doi":"10.1101/290031","title":"Adjutant: an R-based tool to support topic discovery for systematic and literature reviews","year":2018,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Centre for Disease Control; University of British Columbia","funders":"","keywords":"Computer science; Information retrieval; Cluster analysis; Data science; Data mining; Artificial intelligence","score_opus":0.019007988651193183,"score_gpt":0.26826600618695484,"score_spread":0.24925801753576166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2794792323","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045390306,0.0055134525,0.25124684,0.0037820549,0.0013996668,0.0045407107,0.39081553,0.32599205,0.012170757],"genre_scores_gemma":[0.028687784,0.0035668267,0.73230594,0.001618112,0.0005382731,0.021126078,0.1395081,0.06431089,0.008337994],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98667777,0.0059213103,0.0026053358,0.0026603343,0.0017233148,0.0004120807],"domain_scores_gemma":[0.8735701,0.09795194,0.010049322,0.008116974,0.008120408,0.0021912896],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018637525,0.0037567306,0.0048096706,0.021262271,0.0017159579,0.0073398934,0.004191377,0.001671125,0.14809811],"category_scores_gemma":[0.12848501,0.0029425593,0.0073166597,0.018742984,0.0013768544,0.0040639937,0.0075931456,0.0027756088,0.060836684],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017339815,0.000099281606,0.005719753,0.0734892,0.004525562,0.0013331177,0.0018631208,0.004177049,0.007247542,0.018949294,0.7464683,0.13439381],"study_design_scores_gemma":[0.0025942419,0.00033894705,0.008029895,0.0088154115,0.0046019224,0.0019920927,0.00047928872,0.024661226,0.010251057,0.06622584,0.87137604,0.00063395547],"about_ca_topic_score_codex":0.0020369445,"about_ca_topic_score_gemma":0.005602484,"teacher_disagreement_score":0.98136246,"about_ca_system_score_codex":0.0021459535,"about_ca_system_score_gemma":0.013278205,"threshold_uncertainty_score":0.49543756},"labels":[],"label_agreement":null},{"id":"W2796289399","doi":"10.1016/j.artmed.2018.03.004","title":"A two-step approach for mining patient treatment pathways in administrative healthcare databases","year":2018,"lang":"en","type":"article","venue":"Artificial Intelligence in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":38,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre intégré de santé et de services sociaux de Chaudière-Appalaches; Université Laval","funders":"Canadian Institutes of Health Research","keywords":"Cluster analysis; Construct (python library); Medical record; Computer science; Health care; Homogeneous; Cluster (spacecraft); Knowledge extraction; Data science; Data mining; Database; Medicine; Artificial intelligence","score_opus":0.22683987597058422,"score_gpt":0.4223980709586187,"score_spread":0.19555819498803448,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2796289399","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024064032,0.00057410664,0.9563018,0.0018100472,0.00010778568,0.0020296571,0.008786982,0.0042017223,0.002123744],"genre_scores_gemma":[0.054839227,0.00021977242,0.93603575,0.00026553564,0.000026039661,0.00043919077,0.006643227,0.00006937046,0.0014618395],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945979,0.0011119867,0.0010588765,0.0011393308,0.0018288245,0.00026296588],"domain_scores_gemma":[0.9912061,0.0050914614,0.0005535788,0.0011180894,0.0017155125,0.00031527146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004455356,0.0011783324,0.0017364706,0.008026878,0.0019450985,0.0059544463,0.0027374781,0.0026276833,0.004297184],"category_scores_gemma":[0.016008794,0.0009792532,0.0041681654,0.0075274305,0.00064307824,0.0039809668,0.0041683004,0.0022927183,0.0023263874],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012566245,0.0021389208,0.040128104,0.002818529,0.001756001,0.003181632,0.0029538972,0.03380882,0.025702354,0.028249217,0.021015484,0.83699036],"study_design_scores_gemma":[0.00035009105,0.00093620893,0.026381087,0.0007877972,0.0019603216,0.006857382,0.0036691169,0.7572183,0.04688323,0.096147604,0.0584516,0.00035730627],"about_ca_topic_score_codex":0.008161809,"about_ca_topic_score_gemma":0.019799085,"teacher_disagreement_score":0.008161809,"about_ca_system_score_codex":0.001350731,"about_ca_system_score_gemma":0.006572473,"threshold_uncertainty_score":0.023562491},"labels":[],"label_agreement":null},{"id":"W2797078590","doi":"10.1093/jamia/ocy021","title":"UMLS to DBPedia link discovery through circular resolution","year":2018,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Unified Medical Language System; Computer science; Information retrieval; Annotation; Set (abstract data type); Simple Knowledge Organization System; Natural language processing; Ontology; Thesaurus; Artificial intelligence; Semantic Web; RDF; SPARQL","score_opus":0.010433604294806196,"score_gpt":0.291618565749867,"score_spread":0.2811849614550608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2797078590","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007386687,0.0008117158,0.9570989,0.0009760852,0.0003802781,0.00044305963,0.00720324,0.016620727,0.00907934],"genre_scores_gemma":[0.04573123,0.00075698976,0.9171794,0.0007810955,0.00011531794,0.00042320363,0.02780839,0.002219431,0.0049849623],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9876715,0.003998433,0.0013182622,0.0030029297,0.0035863996,0.0004225342],"domain_scores_gemma":[0.9795257,0.0073436643,0.0014950025,0.006428531,0.004753102,0.00045408015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010729731,0.0018603994,0.0011090371,0.012053199,0.003137649,0.0066426634,0.0035221498,0.0016634278,0.006526298],"category_scores_gemma":[0.03908367,0.0014964252,0.0028787586,0.009357738,0.0015107032,0.007344742,0.011173941,0.0033625814,0.0060562184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050865207,0.00048307382,0.0070653297,0.0026245902,0.0007922096,0.0020382241,0.004877807,0.027663764,0.012033193,0.101608045,0.17617318,0.66413194],"study_design_scores_gemma":[0.00012473039,0.00009552791,0.0035254434,0.000987405,0.0004446152,0.0012924494,0.0032405464,0.24945459,0.055528812,0.16049545,0.5245087,0.0003018166],"about_ca_topic_score_codex":0.019135382,"about_ca_topic_score_gemma":0.020416338,"teacher_disagreement_score":0.019135382,"about_ca_system_score_codex":0.0017776773,"about_ca_system_score_gemma":0.0057005994,"threshold_uncertainty_score":0.056744933},"labels":[],"label_agreement":null},{"id":"W2799289533","doi":"","title":"Construction of a data link system for integrating genomic information across species.","year":2017,"lang":"en","type":"article","venue":"The Japanese Biochemical Society/The Molecular Biology Society of Japan","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Genome Canada","funders":"","keywords":"Link (geometry); Computer science; Genomics; Computational biology; Biology; Data science; Genome; Genetics; Gene; Computer network","score_opus":0.029738732840891346,"score_gpt":0.31781842813619987,"score_spread":0.2880796952953085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2799289533","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03966149,0.00046807967,0.85775024,0.00089785585,0.00036255247,0.0015468233,0.035712704,0.057178717,0.0064215437],"genre_scores_gemma":[0.08001985,0.00039196975,0.80019397,0.000273316,0.000064675536,0.0015345754,0.11192755,0.0021055685,0.0034884878],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9984158,0.0003048885,0.0003065337,0.00040425078,0.00047526427,0.0000933114],"domain_scores_gemma":[0.9959417,0.0013932167,0.0002926841,0.0008188162,0.0011985111,0.00035499674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042551323,0.0007509975,0.0007800444,0.0066215782,0.001939904,0.002870832,0.0016787482,0.0011485158,0.004978394],"category_scores_gemma":[0.0106917005,0.00070691673,0.0012948125,0.006035962,0.0005636117,0.0033390264,0.003580323,0.0013599376,0.004011059],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014349082,0.001350178,0.03461344,0.0027269353,0.0007224123,0.0026310352,0.0024448368,0.015941307,0.10308601,0.05401741,0.13561793,0.64541364],"study_design_scores_gemma":[0.00037631183,0.00067377713,0.026336623,0.00089964486,0.0011925674,0.0014864958,0.0020696954,0.2752884,0.19297428,0.068156086,0.43017355,0.00037257196],"about_ca_topic_score_codex":0.004887105,"about_ca_topic_score_gemma":0.0039527817,"teacher_disagreement_score":0.0066215782,"about_ca_system_score_codex":0.0009225115,"about_ca_system_score_gemma":0.004147748,"threshold_uncertainty_score":0.022503555},"labels":[],"label_agreement":null},{"id":"W2799890670","doi":"10.2196/publichealth.8552","title":"Trends in HIV Terminology: Text Mining and Data Visualization Assessment of International AIDS Conference Abstracts Over 25 Years","year":2018,"lang":"en","type":"article","venue":"JMIR Public Health and Surveillance","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"International AIDS Society","keywords":"Terminology; Human immunodeficiency virus (HIV); Data science; Computer science; Data visualization; Visualization; Information retrieval; Medicine; Data mining; Family medicine; Linguistics","score_opus":0.06299404253400136,"score_gpt":0.4072134443898986,"score_spread":0.3442194018558972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2799890670","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.79623437,0.009864264,0.005096118,0.0023424095,0.00045002962,0.0011174557,0.16999461,0.0011296973,0.01377107],"genre_scores_gemma":[0.8114675,0.0056384695,0.035214063,0.00034180234,0.0004945083,0.0029959984,0.13794321,0.00030224695,0.00560222],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9933513,0.0013510411,0.0015767135,0.0009411606,0.0024719615,0.000307878],"domain_scores_gemma":[0.950733,0.02231483,0.012230619,0.0015135767,0.011838924,0.0013691982],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0067532905,0.00066377735,0.00059657864,0.026276778,0.000808243,0.003727835,0.0008606165,0.0005809538,0.0025373963],"category_scores_gemma":[0.04100932,0.000200147,0.0008894997,0.023277804,0.0004634035,0.0025425667,0.0022252854,0.00068314775,0.00114052],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009015591,0.00022974971,0.5042218,0.009379514,0.00039661792,0.0013408668,0.021725275,0.0031643598,0.007956777,0.0022802802,0.067806855,0.38059634],"study_design_scores_gemma":[0.00003774551,0.00038021425,0.80758333,0.001964648,0.00022106718,0.0011709128,0.02057926,0.006104952,0.0057329424,0.0019859194,0.15406457,0.00017446342],"about_ca_topic_score_codex":0.003948384,"about_ca_topic_score_gemma":0.005951464,"teacher_disagreement_score":0.99324673,"about_ca_system_score_codex":0.0013253165,"about_ca_system_score_gemma":0.0016741643,"threshold_uncertainty_score":0.035715282},"labels":[],"label_agreement":null},{"id":"W2800094508","doi":"10.1200/jco.2017.35.15_suppl.6501","title":"Cognitive technology addressing optimal cancer clinical trial matching and protocol feasibility in a community cancer practice.","year":2017,"lang":"en","type":"article","venue":"Journal of Clinical Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"IBM (Canada)","funders":"","keywords":"Medicine; Protocol (science); Population; Clinical trial; Inclusion and exclusion criteria; Breast cancer; IBM; Randomized controlled trial; Cancer; Medical physics; Family medicine; Internal medicine; Alternative medicine; Pathology","score_opus":0.4296981489753393,"score_gpt":0.6408233170785116,"score_spread":0.21112516810317228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2800094508","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.756826,0.0029609192,0.13425998,0.038122714,0.00035299102,0.012546837,0.00055030826,0.0038333044,0.050546866],"genre_scores_gemma":[0.74448586,0.0006023704,0.24669047,0.002498845,0.00013219242,0.0035501865,0.00018286152,0.00013711874,0.0017201751],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9337388,0.054697808,0.0026085563,0.0028767334,0.004614152,0.0014638971],"domain_scores_gemma":[0.6865638,0.27943462,0.010983856,0.0062527237,0.009238791,0.0075261374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.093747206,0.0006293284,0.00065992336,0.0033735796,0.0029729344,0.006224645,0.0025106715,0.0019326112,0.010540802],"category_scores_gemma":[0.3130857,0.00081136165,0.0009589161,0.002168121,0.0020042716,0.0043398016,0.006816729,0.0017409225,0.0011779565],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034840137,0.0043801405,0.039432604,0.0018480233,0.00028185424,0.0005658287,0.015241853,0.0057435976,0.0023514784,0.01071512,0.020316139,0.8956393],"study_design_scores_gemma":[0.014580411,0.02612034,0.16178812,0.0051170927,0.0021070861,0.0034689838,0.06968535,0.34097904,0.017341742,0.21463497,0.1430699,0.0011069044],"about_ca_topic_score_codex":0.007863504,"about_ca_topic_score_gemma":0.009970872,"teacher_disagreement_score":0.093747206,"about_ca_system_score_codex":0.0067596077,"about_ca_system_score_gemma":0.02384114,"threshold_uncertainty_score":0.4957885},"labels":[],"label_agreement":null},{"id":"W2801170057","doi":"10.2196/10281","title":"A Deep Learning Method to Automatically Identify Reports of Scientifically Rigorous Clinical Research from the Biomedical Literature: Comparative Analytic Study","year":2018,"lang":"en","type":"article","venue":"Journal of Medical Internet Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":55,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Impact","funders":"U.S. National Library of Medicine; National Institute of Diabetes and Digestive and Kidney Diseases; Advanced Research Projects Agency; National Cancer Institute; Canadian Institutes of Health Research; Defense Advanced Research Projects Agency","keywords":"Computer science; Data science; Artificial intelligence; Management science; Information retrieval; Engineering","score_opus":0.20439039664998043,"score_gpt":0.5821262361261871,"score_spread":0.37773583947620665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2801170057","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4572658,0.13545498,0.35338363,0.009464724,0.0008735387,0.0028047129,0.027388856,0.0034295747,0.009934178],"genre_scores_gemma":[0.8241807,0.0075265323,0.15115483,0.0013660026,0.0003258625,0.0013896328,0.012887744,0.00010633828,0.0010623545],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98530257,0.008405208,0.0021335918,0.0016818015,0.0022616116,0.0002152166],"domain_scores_gemma":[0.8828376,0.09865666,0.006934852,0.005049017,0.005752364,0.0007695844],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03029847,0.0010060506,0.0013354562,0.01075421,0.00046074265,0.0019452058,0.0016272855,0.0015212232,0.0033782192],"category_scores_gemma":[0.11325833,0.0004465741,0.0023108886,0.0044004745,0.00092359807,0.0023409307,0.002142306,0.0013870397,0.0005651232],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063120406,0.0009293444,0.09442398,0.0139936935,0.007234347,0.00038660417,0.00043607247,0.023116719,0.0034725147,0.004930339,0.014835092,0.8299293],"study_design_scores_gemma":[0.004646169,0.005363147,0.10507754,0.008625081,0.014080461,0.0026940915,0.0009756803,0.73628813,0.020287367,0.046744578,0.054819066,0.00039867213],"about_ca_topic_score_codex":0.004536351,"about_ca_topic_score_gemma":0.007116871,"teacher_disagreement_score":0.9697015,"about_ca_system_score_codex":0.0024771788,"about_ca_system_score_gemma":0.0048709316,"threshold_uncertainty_score":0.16023552},"labels":[],"label_agreement":null},{"id":"W2802392436","doi":"10.3389/fphar.2018.00435","title":"A Semantic Transformation Methodology for the Secondary Use of Observational Healthcare Data in Postmarketing Safety Studies","year":2018,"lang":"en","type":"article","venue":"Frontiers in Pharmacology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"Seventh Framework Programme","keywords":"Observational study; Health care; Medicine; Postmarketing surveillance; Patient safety; Data science; Computer science; Adverse effect; Pharmacology; Risk analysis (engineering); Internal medicine","score_opus":0.3042727212310511,"score_gpt":0.44921670265698094,"score_spread":0.14494398142592985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2802392436","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018219551,0.00004723322,0.99486625,0.00026396808,0.000029360011,0.00033599985,0.0004264951,0.0010164108,0.0011923977],"genre_scores_gemma":[0.02970674,0.00014362579,0.96606964,0.00018671052,0.000032980886,0.00059772667,0.002293799,0.00026478848,0.0007038799],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9825827,0.0067294487,0.0030246642,0.0022241005,0.005023573,0.00041560174],"domain_scores_gemma":[0.9833745,0.0067636813,0.0014080835,0.0044273967,0.0036206814,0.00040564712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022475107,0.00088294875,0.0007552074,0.008220059,0.0014480809,0.005129588,0.0023738716,0.0014703956,0.0025956316],"category_scores_gemma":[0.022878373,0.00084590376,0.004226762,0.0071016015,0.0025998657,0.007569228,0.0075487234,0.0029065679,0.0013448271],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018622757,0.00052450836,0.0067657824,0.0012070702,0.0003782041,0.0013364622,0.0061891903,0.019917881,0.014261149,0.59991217,0.009467276,0.33985406],"study_design_scores_gemma":[0.00012131181,0.00023280596,0.006793156,0.0009850272,0.00034647388,0.0015804995,0.0039539966,0.1921026,0.040940505,0.50066006,0.25203058,0.00025300268],"about_ca_topic_score_codex":0.0036932493,"about_ca_topic_score_gemma":0.0044225543,"teacher_disagreement_score":0.022475107,"about_ca_system_score_codex":0.0020012655,"about_ca_system_score_gemma":0.0072044106,"threshold_uncertainty_score":0.11886114},"labels":[],"label_agreement":null},{"id":"W2804447071","doi":"10.5771/9783956504211-392","title":"Photography as a legitimate technique for domain analysis in Knowledge Organization","year":2018,"lang":"en","type":"book-chapter","venue":"Ergon Verlag eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Photography; Domain (mathematical analysis); Domain knowledge; Computer science; Knowledge management; Visual arts; Art; Mathematics","score_opus":0.011040308625992916,"score_gpt":0.26378598282626126,"score_spread":0.25274567420026833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2804447071","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13813603,0.007694237,0.36355704,0.015271027,0.0022499694,0.0012738832,0.0028434615,0.0019849036,0.4669895],"genre_scores_gemma":[0.64876515,0.0032926763,0.3206235,0.0018581864,0.00043577317,0.0017762808,0.00065640936,0.0009679592,0.021624018],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9934052,0.0043503637,0.0002652847,0.0006233141,0.0010832398,0.00027262513],"domain_scores_gemma":[0.9760275,0.01762797,0.0012988813,0.0031106598,0.001573134,0.00036182906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008306351,0.0006071644,0.00034707753,0.008454017,0.0036134664,0.007125418,0.0009359154,0.0013291637,0.01624563],"category_scores_gemma":[0.021800056,0.0004461463,0.00065648986,0.0059692594,0.006180193,0.0070653376,0.005260141,0.002987102,0.001899483],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029896724,0.00013573618,0.013005145,0.0029305343,0.000084129206,0.0011320051,0.2671358,0.0011672336,0.014520845,0.32745197,0.051107142,0.3210305],"study_design_scores_gemma":[0.00007345844,0.00013582452,0.029539565,0.0025520553,0.00008846366,0.0018157564,0.15090147,0.0040840786,0.0088905245,0.12698616,0.67479247,0.00014023941],"about_ca_topic_score_codex":0.003925448,"about_ca_topic_score_gemma":0.008183287,"teacher_disagreement_score":0.01624563,"about_ca_system_score_codex":0.0020905402,"about_ca_system_score_gemma":0.0019231947,"threshold_uncertainty_score":0.05434704},"labels":[],"label_agreement":null},{"id":"W2805089341","doi":"","title":"IBM Research System at TAC 2017: Adverse Drug Reactions Extraction from Drug Labels.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"IBM; Drug; Computer science; Adverse drug reaction; Drug reaction; Extraction (chemistry); Pharmacology; Medicine; Chemistry; Chromatography; Nanotechnology; Materials science","score_opus":0.026186522099544664,"score_gpt":0.34179864953474665,"score_spread":0.315612127435202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805089341","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011123589,0.002430216,0.07709188,0.0015556534,0.0005428754,0.0008639165,0.5497558,0.34350538,0.013130662],"genre_scores_gemma":[0.038836226,0.0012458256,0.149765,0.00057378743,0.00028935252,0.0010457281,0.792952,0.0069876797,0.008304387],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982716,0.00034213805,0.0002710743,0.00049857766,0.00048763,0.00012887434],"domain_scores_gemma":[0.9961635,0.0012401366,0.00046764076,0.00085997535,0.0009994338,0.00026938278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002574354,0.0023537804,0.0019908492,0.007353097,0.00089005457,0.003009595,0.0018878847,0.0016354129,0.030952431],"category_scores_gemma":[0.012301782,0.00079174805,0.0016234067,0.005535023,0.00038455613,0.0030574242,0.0022460138,0.0016124209,0.038264252],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001065523,0.0002772394,0.004045145,0.0020573223,0.0003945646,0.00028892973,0.00021306063,0.002990477,0.00603396,0.004167015,0.81548834,0.16297844],"study_design_scores_gemma":[0.001395083,0.00063771405,0.016344294,0.0007785513,0.00081144465,0.000878665,0.00042219923,0.14150931,0.030293843,0.046199467,0.7604485,0.00028098447],"about_ca_topic_score_codex":0.0096331015,"about_ca_topic_score_gemma":0.008884442,"teacher_disagreement_score":0.030952431,"about_ca_system_score_codex":0.0012249746,"about_ca_system_score_gemma":0.0035220408,"threshold_uncertainty_score":0.1035462},"labels":[],"label_agreement":null},{"id":"W2805208499","doi":"","title":"SPARQL Assist Language-Neutral Query Composer.","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital; University of British Columbia","funders":"","keywords":"SPARQL; Computer science; Identifier; Task (project management); Query language; Named graph; Information retrieval; Composition (language); World Wide Web; Natural language processing; Linguistics; Semantic Web; Programming language; RDF","score_opus":0.0059585685757162915,"score_gpt":0.270039079173995,"score_spread":0.2640805105982787,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805208499","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0067640366,0.00048656142,0.6831341,0.0012728715,0.00032463577,0.00094249245,0.008805375,0.27773538,0.020534517],"genre_scores_gemma":[0.17882197,0.00094798935,0.66927093,0.004314619,0.0004620581,0.0012412232,0.044494793,0.05301266,0.04743365],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9938123,0.0015402654,0.0007957544,0.0012331663,0.0020961568,0.0005223644],"domain_scores_gemma":[0.9914659,0.003288484,0.00032239468,0.0025671273,0.0019654278,0.00039059142],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0084492015,0.0015040538,0.00095783995,0.0010302317,0.0009049864,0.0035815157,0.0032803735,0.0014085786,0.05036239],"category_scores_gemma":[0.016921083,0.000913367,0.0012428559,0.0011397743,0.0009962883,0.005275876,0.0045961626,0.0019837408,0.027230065],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026145612,0.00046568652,0.005066082,0.00215305,0.00025498346,0.0016685593,0.002431521,0.0049571744,0.050240338,0.088159434,0.4057258,0.43626282],"study_design_scores_gemma":[0.00033214188,0.00018604544,0.0015316128,0.00034158677,0.00013027813,0.0020277323,0.00072461716,0.08560536,0.0940646,0.074132405,0.7407284,0.00019524449],"about_ca_topic_score_codex":0.0023037177,"about_ca_topic_score_gemma":0.001980005,"teacher_disagreement_score":0.05036239,"about_ca_system_score_codex":0.0010483189,"about_ca_system_score_gemma":0.0015452867,"threshold_uncertainty_score":0.16847897},"labels":[],"label_agreement":null},{"id":"W2805893819","doi":"10.1016/j.jbi.2018.06.001","title":"Patient similarity for precision medicine: A systematic review","year":2018,"lang":"en","type":"review","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":163,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Precision medicine; Computer science; Profiling (computer programming); Psychosocial; Data mining; Data science; Similarity (geometry); Identification (biology); Medicine; Cluster analysis; Data integration; Systematic review; Guideline; MEDLINE; Artificial intelligence; Machine learning; Pathology","score_opus":0.04857467617237364,"score_gpt":0.37806832901913123,"score_spread":0.3294936528467576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2805893819","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00041735603,0.9984497,0.00026072617,0.0003058208,0.00009829709,0.000106465886,0.00022633668,0.000006365816,0.00012898057],"genre_scores_gemma":[0.016666792,0.9786453,0.0025118804,0.0013024476,0.00015844611,0.00023570203,0.0003903942,0.000009607131,0.00007947835],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9814844,0.006515767,0.0068519497,0.0017403563,0.0030931286,0.00031425827],"domain_scores_gemma":[0.9335193,0.053207304,0.008174846,0.0013662032,0.0032545987,0.00047774578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.018304326,0.0017102341,0.011640594,0.014790791,0.00075912924,0.004050943,0.0032229836,0.0028489179,0.005215495],"category_scores_gemma":[0.07479848,0.0012536645,0.009437712,0.012289448,0.0013793835,0.005265079,0.003175703,0.0021783174,0.00037623505],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003989224,0.000044398643,0.0018595384,0.84355134,0.02845849,0.00012471373,0.0002561616,0.00021918834,0.00021541532,0.00085275806,0.0039578835,0.120061085],"study_design_scores_gemma":[0.0009287392,0.00036530165,0.0071750535,0.7521265,0.19395256,0.00088704465,0.0005319756,0.00038711773,0.00039819977,0.0024126081,0.040684607,0.0001501473],"about_ca_topic_score_codex":0.0037588154,"about_ca_topic_score_gemma":0.013622816,"teacher_disagreement_score":0.018304326,"about_ca_system_score_codex":0.0029227564,"about_ca_system_score_gemma":0.011919378,"threshold_uncertainty_score":0.096803725},"labels":[],"label_agreement":null},{"id":"W2806036031","doi":"","title":"Extracting and Normalizing Adverse Drug Reactions from Drug Labels.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug; Drug reaction; Computer science; Adverse drug reaction; Pharmacology; Medicine","score_opus":0.010564392424887884,"score_gpt":0.2759563180880878,"score_spread":0.2653919256631999,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806036031","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12823744,0.009504272,0.70504284,0.0028178184,0.001282177,0.002163388,0.10516707,0.016554836,0.029230203],"genre_scores_gemma":[0.2805107,0.003409535,0.5901183,0.00053257163,0.0003819073,0.0012201923,0.11563102,0.0008208453,0.007374874],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.996635,0.00062996795,0.00054728653,0.0007592475,0.0012160907,0.00021237812],"domain_scores_gemma":[0.9916728,0.003386541,0.0012892748,0.0009292455,0.00250406,0.00021809695],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002642212,0.0008018864,0.0008593341,0.012598283,0.0008317865,0.0019838721,0.0010746205,0.00086089125,0.0038169299],"category_scores_gemma":[0.016011167,0.00025286255,0.0012991856,0.008004058,0.0005099322,0.0016045284,0.0016333092,0.0010241654,0.003393018],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005148573,0.0003214302,0.035987888,0.0022666159,0.000374152,0.0007492545,0.00065910444,0.0031858333,0.018942423,0.013970156,0.056648947,0.86637944],"study_design_scores_gemma":[0.00024622225,0.0005066656,0.12894249,0.0017957127,0.001637102,0.004725155,0.0039764266,0.12405011,0.10025453,0.15228139,0.48127756,0.0003066241],"about_ca_topic_score_codex":0.0062943324,"about_ca_topic_score_gemma":0.009753747,"teacher_disagreement_score":0.012598283,"about_ca_system_score_codex":0.00090163486,"about_ca_system_score_gemma":0.0037915134,"threshold_uncertainty_score":0.013973534},"labels":[],"label_agreement":null},{"id":"W2806165888","doi":"","title":"Extracting Adverse Drug Reactions using Deep Learning and Dictionary Based Approaches.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Drug reaction; Artificial intelligence; Dictionary learning; Natural language processing; Drug; Pharmacology; Medicine; Sparse approximation","score_opus":0.023596337213031292,"score_gpt":0.28534421284367584,"score_spread":0.26174787563064456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806165888","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22235316,0.02019391,0.66267806,0.004152291,0.001135979,0.0015546461,0.056224704,0.0068276883,0.024879556],"genre_scores_gemma":[0.60746396,0.0066047073,0.31875533,0.0008458945,0.00033693,0.0005776327,0.058628768,0.00018398586,0.006602862],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99874073,0.00024109264,0.00024736283,0.0002269327,0.00042885443,0.00011501179],"domain_scores_gemma":[0.9965013,0.0016419282,0.0005952682,0.00031561585,0.0008236963,0.00012215931],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00085552194,0.0006001196,0.0008772856,0.005737559,0.00041594676,0.0010318413,0.0008008643,0.00086853286,0.002462433],"category_scores_gemma":[0.00550253,0.00018921988,0.0010041317,0.003779343,0.00027361678,0.0013689774,0.0013803836,0.00089765445,0.001500925],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00053590466,0.00047545315,0.030406157,0.002021466,0.00042960944,0.0011156992,0.00022649863,0.009876679,0.013412557,0.0069369418,0.039827365,0.8947357],"study_design_scores_gemma":[0.00028501436,0.0009849101,0.07133677,0.0014715584,0.0013958083,0.005140274,0.0019820367,0.5644564,0.04541536,0.1308341,0.17648233,0.00021533637],"about_ca_topic_score_codex":0.004717498,"about_ca_topic_score_gemma":0.008980292,"teacher_disagreement_score":0.005737559,"about_ca_system_score_codex":0.0006506872,"about_ca_system_score_gemma":0.0016334249,"threshold_uncertainty_score":0.009380102},"labels":[],"label_agreement":null},{"id":"W2806398775","doi":"10.2196/medinform.7799","title":"Development and Validation of a Functional Behavioural Assessment Ontology to Support Behavioural Health Interventions","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Psychological intervention; Ontology; Computer science; Psychology; Medicine; Nursing","score_opus":0.07670011050300879,"score_gpt":0.39281468987887525,"score_spread":0.31611457937586646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2806398775","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041141048,0.00021903463,0.9288657,0.0017988257,0.00013457704,0.007542235,0.0023718653,0.003849591,0.014077211],"genre_scores_gemma":[0.08486877,0.00019922123,0.9043175,0.00022206239,0.000017241793,0.0032360908,0.0051393234,0.0003390498,0.001660743],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9874505,0.0048429943,0.002373717,0.001255395,0.0035593226,0.0005181751],"domain_scores_gemma":[0.9626628,0.016020251,0.002722078,0.005273613,0.012474264,0.0008470494],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021012306,0.0010043049,0.0006583309,0.004810118,0.0017719277,0.0036324817,0.00268183,0.0013559678,0.0025360147],"category_scores_gemma":[0.04135814,0.00078647607,0.0028111648,0.0020344094,0.002253533,0.004209463,0.0040517584,0.0021645618,0.0009197548],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036345236,0.0021891452,0.038845424,0.0037366506,0.0003901898,0.0010055557,0.016469205,0.040242508,0.021912707,0.1566762,0.017819636,0.7003492],"study_design_scores_gemma":[0.00040327752,0.0010444603,0.049418114,0.008089885,0.00083620514,0.0016858946,0.009862139,0.29497337,0.042722665,0.14986156,0.44057927,0.0005232579],"about_ca_topic_score_codex":0.017985182,"about_ca_topic_score_gemma":0.014364596,"teacher_disagreement_score":0.021012306,"about_ca_system_score_codex":0.005492876,"about_ca_system_score_gemma":0.0216721,"threshold_uncertainty_score":0.11112505},"labels":[],"label_agreement":null},{"id":"W2807245502","doi":"10.63317/2k4h99c3w2uz","title":"A New Corpus to Support Text Mining for the Curation of Metabolites in the ChEBI Database","year":2018,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"European Commission","keywords":"Computer science; Matching (statistics); Protocol (science); Information retrieval; Data curation; Biomedical text mining; Natural language processing; Text mining; Data mining","score_opus":0.04024246612935126,"score_gpt":0.3308117969133884,"score_spread":0.2905693307840371,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807245502","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06740744,0.008087261,0.102898896,0.0034332855,0.0018176193,0.002366747,0.77477026,0.017634394,0.0215841],"genre_scores_gemma":[0.043096527,0.0023583085,0.22352059,0.00054856006,0.00033728994,0.0032804285,0.71869814,0.0021570844,0.0060030892],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977055,0.0004131215,0.00056591496,0.00061650627,0.00060618465,0.000092813665],"domain_scores_gemma":[0.9881264,0.0061519174,0.0008473848,0.0014383296,0.0026675502,0.000768401],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002813018,0.0011762073,0.0013988647,0.014445985,0.0020549225,0.003029511,0.0011880787,0.0013911964,0.013792089],"category_scores_gemma":[0.01330609,0.000595648,0.0009004359,0.013062636,0.0009283168,0.0036206155,0.0032023236,0.0020468116,0.0075072944],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020534005,0.0009032878,0.011012327,0.018427277,0.00049056974,0.0037669057,0.0039488543,0.0043863314,0.11126068,0.01992368,0.4846738,0.339153],"study_design_scores_gemma":[0.00034408201,0.00018528405,0.02254915,0.0009161888,0.00032231,0.0016225739,0.0011806532,0.010073433,0.023512641,0.00857811,0.9305224,0.00019318322],"about_ca_topic_score_codex":0.004975778,"about_ca_topic_score_gemma":0.00928789,"teacher_disagreement_score":0.014445985,"about_ca_system_score_codex":0.0013027175,"about_ca_system_score_gemma":0.0041431044,"threshold_uncertainty_score":0.04613918},"labels":[],"label_agreement":null},{"id":"W2807391974","doi":"","title":"TCS Research at TAC 2017: Joint Extraction of Entities and Relations from Drug Labels using an Ensemble of Neural Networks.","year":2017,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial neural network; Joint (building); Extraction (chemistry); Artificial intelligence; Machine learning; Chemistry; Chromatography; Engineering","score_opus":0.05898442574939388,"score_gpt":0.3593710565403336,"score_spread":0.3003866307909397,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2807391974","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13892092,0.010764745,0.40285087,0.008033823,0.003816491,0.0019652282,0.2691601,0.13274285,0.031744953],"genre_scores_gemma":[0.15020894,0.00199782,0.45851004,0.000742488,0.00052394555,0.0010016015,0.36618236,0.002743458,0.018089399],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99804026,0.00049431616,0.00012428653,0.00050628255,0.0006594557,0.00017533262],"domain_scores_gemma":[0.99708754,0.0009365871,0.00015467481,0.00065047847,0.0009038383,0.0002669474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027535774,0.0025845128,0.0015716994,0.005277563,0.0013181443,0.0025830434,0.002251032,0.002260244,0.008895466],"category_scores_gemma":[0.0072975014,0.0007000279,0.0020107718,0.004085896,0.0004521624,0.003431906,0.0019478118,0.0032513523,0.0069823195],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012176925,0.0016193298,0.006371639,0.001158429,0.0011599768,0.00061705295,0.00029617414,0.029882034,0.01906748,0.0067728725,0.51565075,0.4161866],"study_design_scores_gemma":[0.0005076424,0.00057853904,0.006305721,0.0002579046,0.0008427018,0.0003659095,0.00045779056,0.7482213,0.03176649,0.032138113,0.17837021,0.00018772375],"about_ca_topic_score_codex":0.025742358,"about_ca_topic_score_gemma":0.04593501,"teacher_disagreement_score":0.025742358,"about_ca_system_score_codex":0.0016799594,"about_ca_system_score_gemma":0.005014729,"threshold_uncertainty_score":0.051185012},"labels":[],"label_agreement":null},{"id":"W2808426020","doi":"10.2196/10205","title":"The D2Refine Platform for the Standardization of Clinical Research Study Data Dictionaries: Usability Study","year":2018,"lang":"en","type":"article","venue":"JMIR Human Factors","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of General Medical Sciences; National Cancer Institute; National Institutes of Health","keywords":"Standardization; Terminology; Usability; Harmonization; Computer science; Data science; World Wide Web; Human–computer interaction","score_opus":0.3790427994484278,"score_gpt":0.5483266059884261,"score_spread":0.1692838065399983,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808426020","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73029095,0.0011776476,0.21422775,0.0029636,0.00032908106,0.027979976,0.0033508935,0.009440838,0.010239325],"genre_scores_gemma":[0.4852524,0.00056172913,0.48027962,0.001089823,0.00011028056,0.025074942,0.0031778724,0.001736219,0.002717068],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.91003186,0.07062149,0.007845775,0.0034938967,0.00698897,0.0010180034],"domain_scores_gemma":[0.74366194,0.19706814,0.007535467,0.027461827,0.01973418,0.0045384476],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13020515,0.0007640612,0.0007890659,0.0025093358,0.0011235704,0.0030987598,0.001416278,0.0011467787,0.0028001426],"category_scores_gemma":[0.15318917,0.00097462686,0.0012666303,0.001412538,0.0010968137,0.0039503644,0.006467104,0.0017187201,0.0010240512],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009604484,0.007784857,0.06757977,0.006855641,0.00048393637,0.0011048431,0.070290454,0.0036849554,0.04316668,0.007087847,0.026446946,0.75590956],"study_design_scores_gemma":[0.010274274,0.028923701,0.31975842,0.0075411475,0.0009762651,0.005331834,0.03559067,0.0618001,0.0870355,0.009287939,0.4314379,0.0020422346],"about_ca_topic_score_codex":0.0010136102,"about_ca_topic_score_gemma":0.0019307644,"teacher_disagreement_score":0.86979485,"about_ca_system_score_codex":0.0014457729,"about_ca_system_score_gemma":0.0045483154,"threshold_uncertainty_score":0.6885989},"labels":[],"label_agreement":null},{"id":"W2808574250","doi":"10.13140/rg.2.1.3769.2406/1","title":"EDDA Study Designs Taxonomy (version 2.0)","year":2016,"lang":"en","type":"article","venue":"D-Scholarship@Pitt (University of Pittsburgh)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Terminology; Taxonomy (biology); Plural; Library science; Subject (documents); Computer science; Information retrieval; Linguistics; Philosophy","score_opus":0.05116250546775349,"score_gpt":0.24891814483894184,"score_spread":0.19775563937118834,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808574250","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032532953,0.0148342345,0.25637633,0.005176091,0.0014948548,0.14211085,0.5323448,0.009361018,0.035048563],"genre_scores_gemma":[0.004679877,0.006473816,0.5426077,0.001723085,0.00021189165,0.3453965,0.090065144,0.0020155685,0.00682645],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.91127133,0.03514062,0.038786374,0.0031717569,0.010260648,0.0013692245],"domain_scores_gemma":[0.75687164,0.16833003,0.019061018,0.015018603,0.038522806,0.002196017],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.10454909,0.0019991086,0.0051642684,0.02355075,0.0021329336,0.007316092,0.005000576,0.0029432815,0.11007696],"category_scores_gemma":[0.23157588,0.0048067407,0.008655761,0.025349826,0.0017149647,0.005523908,0.0074302433,0.0068234573,0.023049142],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012854055,0.00020069038,0.004051215,0.11643798,0.0011270187,0.00017005185,0.005230588,0.0026803135,0.0007440413,0.051906433,0.5039257,0.3122406],"study_design_scores_gemma":[0.00081671297,0.00020494667,0.0037849515,0.020879382,0.0005634441,0.00017904943,0.0008730713,0.0009446096,0.00029026277,0.028576868,0.94275355,0.00013317273],"about_ca_topic_score_codex":0.0075020795,"about_ca_topic_score_gemma":0.011942646,"teacher_disagreement_score":0.8954509,"about_ca_system_score_codex":0.0066508185,"about_ca_system_score_gemma":0.031062402,"threshold_uncertainty_score":0.55291504},"labels":[],"label_agreement":null},{"id":"W2809369748","doi":"10.5430/jha.v7n4p60","title":"Integrating semantic and fuzzy dimensions into electronic medical records: Case of cerebral palsy information system","year":2018,"lang":"en","type":"article","venue":"Journal of Hospital Administration","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Interoperability; Semantic interoperability; Health informatics; Ontology; Knowledge management; Semantics (computer science); Software engineering; Data science; Health care; World Wide Web","score_opus":0.005579544278129308,"score_gpt":0.26050744651796265,"score_spread":0.25492790223983336,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809369748","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67280567,0.006866181,0.14717372,0.044056598,0.00039702666,0.0006399017,0.001976425,0.0007779219,0.12530659],"genre_scores_gemma":[0.9112243,0.002599568,0.0755871,0.0008881489,0.00010980743,0.000108890155,0.0006467776,0.000043143948,0.008792259],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99557096,0.0017673356,0.00045332304,0.0003005363,0.001381893,0.00052589446],"domain_scores_gemma":[0.9957677,0.002631336,0.00029176802,0.00036204993,0.00076069275,0.00018647152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003666531,0.00034990127,0.00046662183,0.0026433154,0.0029088599,0.005014096,0.00118651,0.004025756,0.0025588912],"category_scores_gemma":[0.0076844073,0.00030752807,0.001163654,0.004419209,0.0023293365,0.005485088,0.0027284615,0.0018149958,0.00042369042],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069673033,0.00070701283,0.08181511,0.0016383161,0.0002143582,0.09898579,0.046092633,0.040134188,0.007056431,0.48187503,0.029912958,0.21087149],"study_design_scores_gemma":[0.00016249204,0.00038343514,0.039183017,0.0013161667,0.00046114303,0.047922052,0.065406315,0.23068552,0.019581897,0.17120077,0.42336756,0.00032961502],"about_ca_topic_score_codex":0.034485757,"about_ca_topic_score_gemma":0.027048578,"teacher_disagreement_score":0.034485757,"about_ca_system_score_codex":0.0043773465,"about_ca_system_score_gemma":0.0032373986,"threshold_uncertainty_score":0.06857008},"labels":[],"label_agreement":null},{"id":"W2809924027","doi":"10.1093/jamia/ocy074","title":"A systematic assessment of the availability and clinical drug information coverage of machine-readable clinical drug data sources for building knowledge translation products","year":2018,"lang":"en","type":"review","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Memorial University of Newfoundland","funders":"Canadian Institutes of Health Research; Memorial University of Newfoundland; Diabetes Canada","keywords":"Computer science; Drug; Information retrieval; Data source; Information source (mathematics); Knowledge translation; Quality (philosophy); Medicine; Knowledge management; Pharmacology","score_opus":0.07106297838884242,"score_gpt":0.4324555816785575,"score_spread":0.3613926032897151,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809924027","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16991395,0.75733,0.013868561,0.006560697,0.00063512917,0.0077639935,0.0352371,0.0003692137,0.0083213905],"genre_scores_gemma":[0.6939887,0.19684169,0.0679181,0.0025798036,0.00032902736,0.012884837,0.024708616,0.00021088621,0.00053843047],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.80588,0.068097025,0.08433777,0.0074357027,0.033161584,0.0010879264],"domain_scores_gemma":[0.24594533,0.5887712,0.09388058,0.017613797,0.052023835,0.0017652761],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1660717,0.0010297294,0.004729469,0.07635746,0.0014128314,0.0058456417,0.0033282596,0.00235592,0.0021982484],"category_scores_gemma":[0.5018284,0.0013980502,0.004944848,0.04066631,0.002810975,0.008450396,0.0067269676,0.0010981587,0.00038819725],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015123171,0.00017081405,0.077130266,0.574831,0.013493618,0.00073870487,0.0077076475,0.0007044531,0.0018126107,0.0020794272,0.0074783172,0.3123409],"study_design_scores_gemma":[0.00075463305,0.0007626868,0.11798416,0.7574155,0.04461031,0.0017902973,0.004779914,0.0021327685,0.0038883486,0.0017624884,0.06386676,0.00025222846],"about_ca_topic_score_codex":0.0045545353,"about_ca_topic_score_gemma":0.012627743,"teacher_disagreement_score":0.8339283,"about_ca_system_score_codex":0.004413328,"about_ca_system_score_gemma":0.0154416375,"threshold_uncertainty_score":0.8782816},"labels":[],"label_agreement":null},{"id":"W284770367","doi":"10.7202/1033220ar","title":"Nommage de documents électroniques : mise au point et évaluation d’une procédure","year":2015,"lang":"fr","type":"article","venue":"Documentation et bibliothèques","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Humanities; Philosophy; Gynecology; Medicine","score_opus":0.04076340440690202,"score_gpt":0.3674587671242661,"score_spread":0.3266953627173641,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W284770367","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13942418,0.0026765014,0.81994265,0.00282496,0.0005560554,0.0014637582,0.0026006638,0.017172234,0.013338931],"genre_scores_gemma":[0.2221758,0.0010337976,0.7620483,0.00024096912,0.000090434456,0.00053794024,0.0015920579,0.0011451558,0.011135597],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9929028,0.0023473334,0.00089689007,0.0011741508,0.0023882035,0.00029058196],"domain_scores_gemma":[0.96605533,0.016278123,0.0018279033,0.005508707,0.009771972,0.00055794144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011463295,0.0007838517,0.00080393685,0.0051328475,0.002295366,0.0055016847,0.0013328798,0.0014226469,0.0054584797],"category_scores_gemma":[0.045620758,0.00055623543,0.00097127346,0.0038089752,0.0012093597,0.0035805134,0.0020071322,0.0011769554,0.00378572],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013005645,0.00024023883,0.0115375975,0.0014517385,0.0001795968,0.00029471953,0.0040454315,0.0033231573,0.06659219,0.011610829,0.009299703,0.8901242],"study_design_scores_gemma":[0.00037710014,0.0012789804,0.061936382,0.0011591798,0.0006242509,0.0025511067,0.009682799,0.09238468,0.55131775,0.02820249,0.24994783,0.0005374407],"about_ca_topic_score_codex":0.008534895,"about_ca_topic_score_gemma":0.007982744,"teacher_disagreement_score":0.011463295,"about_ca_system_score_codex":0.0014254887,"about_ca_system_score_gemma":0.0033060529,"threshold_uncertainty_score":0.06062442},"labels":[],"label_agreement":null},{"id":"W2883972870","doi":"10.2196/medinform.9979","title":"Identifying Principles for the Construction of an Ontology-Based Knowledge Base: A Case Study Approach","year":2018,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Ontology; Knowledge base; Knowledge management; Base (topology); Software engineering; Data science; World Wide Web; Epistemology","score_opus":0.06751049979267319,"score_gpt":0.37157371072944134,"score_spread":0.30406321093676814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2883972870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15634064,0.0007021689,0.7558876,0.008974478,0.00010183574,0.0030058257,0.00061326084,0.0005362077,0.073837936],"genre_scores_gemma":[0.2795679,0.00052229344,0.7133146,0.00038623062,0.000012159064,0.001035506,0.0005629452,0.00013626009,0.0044621993],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98410255,0.010150734,0.0010860185,0.0007493403,0.0031178354,0.0007935372],"domain_scores_gemma":[0.97967553,0.015343008,0.0007163075,0.0015710699,0.002168615,0.0005254868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021334851,0.00074056984,0.00049716845,0.005994846,0.005894163,0.008668951,0.0032999937,0.0036982866,0.003991531],"category_scores_gemma":[0.02003752,0.0008556189,0.0015912706,0.0046366868,0.0065128994,0.011862708,0.0080910055,0.0032660349,0.000765661],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002121555,0.001322284,0.01233156,0.0015006042,0.0000995442,0.021474317,0.12601201,0.010126282,0.0066652354,0.6309291,0.009751511,0.17957526],"study_design_scores_gemma":[0.00026534474,0.0004703709,0.0048971116,0.0030710173,0.00029055474,0.0172217,0.16196647,0.07772555,0.0204831,0.32116857,0.39216447,0.00027569867],"about_ca_topic_score_codex":0.009968025,"about_ca_topic_score_gemma":0.014536597,"teacher_disagreement_score":0.021334851,"about_ca_system_score_codex":0.0063938927,"about_ca_system_score_gemma":0.008438989,"threshold_uncertainty_score":0.11283088},"labels":[],"label_agreement":null},{"id":"W2885593556","doi":"10.1002/jrsm.1314","title":"Design and implementation of a tool for conversion of search strategies between PubMed and Ovid MEDLINE","year":2018,"lang":"en","type":"article","venue":"Research Synthesis Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College; College of Physicians and Surgeons of Ontario","funders":"","keywords":"MEDLINE; Computer science; Information retrieval; Online search; Syntax; Interface (matter); World Wide Web; Artificial intelligence","score_opus":0.20238052298318512,"score_gpt":0.5209453212250386,"score_spread":0.31856479824185346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885593556","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012517718,0.0007449393,0.6917875,0.0025544362,0.0008556388,0.018747753,0.017487168,0.2453944,0.009910493],"genre_scores_gemma":[0.020568969,0.0003497575,0.9358422,0.0009048655,0.00011793856,0.013860685,0.01007711,0.012129126,0.0061492953],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9808029,0.0062384685,0.006549482,0.0031599118,0.0025844742,0.0006648767],"domain_scores_gemma":[0.8773893,0.08921896,0.0059281243,0.009460273,0.015373759,0.0026295555],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.043352943,0.0026667707,0.0016829362,0.013709245,0.0015362222,0.0071913432,0.0045117177,0.0025869764,0.038926687],"category_scores_gemma":[0.119355306,0.003013081,0.0028567065,0.00746615,0.0013759925,0.0071592457,0.0057911463,0.0028886276,0.015771285],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042667626,0.0012514815,0.006958503,0.013137095,0.0009457893,0.0026140874,0.010053521,0.0033731437,0.025093699,0.0140958885,0.1368062,0.7814037],"study_design_scores_gemma":[0.004542131,0.0022864898,0.012144341,0.006287655,0.0012824247,0.00346397,0.0040238127,0.046617426,0.082058825,0.01852931,0.81727964,0.0014840041],"about_ca_topic_score_codex":0.0033652578,"about_ca_topic_score_gemma":0.0039584124,"teacher_disagreement_score":0.95664704,"about_ca_system_score_codex":0.002472715,"about_ca_system_score_gemma":0.007998911,"threshold_uncertainty_score":0.22927505},"labels":[],"label_agreement":null},{"id":"W2885913358","doi":"","title":"What is a risk? A formal representation of risk of stroke for people with atrial fibrillation","year":2017,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Atrial fibrillation; Stroke risk; Stroke (engine); Representation (politics); Computer science; Medicine; Cardiology; Ischemic stroke; Engineering; Political science","score_opus":0.016812385671220627,"score_gpt":0.2756471909601792,"score_spread":0.2588348052889586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2885913358","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.102551684,0.0020715187,0.833086,0.015616217,0.0005568211,0.0003987711,0.019535607,0.0042556575,0.021927763],"genre_scores_gemma":[0.6997804,0.0010054244,0.28386912,0.0009202627,0.00016953549,0.00030437534,0.0098224105,0.00017603986,0.003952492],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9982956,0.0005570116,0.0003360487,0.00033953373,0.00030776713,0.00016408361],"domain_scores_gemma":[0.99478585,0.0031911698,0.0006113872,0.00045259204,0.00064377056,0.00031520656],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025159793,0.00073416246,0.00038653976,0.0027536782,0.0008074047,0.003295375,0.0009344046,0.0012954886,0.004125755],"category_scores_gemma":[0.011004377,0.000315782,0.0016152398,0.0019042786,0.0009956601,0.004766528,0.0016302065,0.0012929708,0.00050823187],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046377236,0.00026076476,0.032942127,0.00097569235,0.00038434463,0.002344514,0.0069585303,0.049570095,0.002758374,0.75947344,0.022111759,0.121756606],"study_design_scores_gemma":[0.00010033619,0.00014609107,0.007763643,0.0007535232,0.00070250256,0.0025429581,0.0025139034,0.21013296,0.002184663,0.65661675,0.11641503,0.00012766312],"about_ca_topic_score_codex":0.016136354,"about_ca_topic_score_gemma":0.016837012,"teacher_disagreement_score":0.016136354,"about_ca_system_score_codex":0.0015405472,"about_ca_system_score_gemma":0.001857256,"threshold_uncertainty_score":0.032084823},"labels":[],"label_agreement":null},{"id":"W2886305294","doi":"10.1109/civemsa.2018.8439958","title":"Pictorial Visualization of EMR Summary Interface and Medical Information Extraction of Clinical Notes","year":2018,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; University of Ottawa","funders":"","keywords":"Computer science; Timeline; Information extraction; Visualization; Interface (matter); Information retrieval; User interface; Graphical user interface; Representation (politics); Human–computer interaction; Information visualization; Artificial intelligence; Natural language processing","score_opus":0.02728687747125595,"score_gpt":0.40152453386047454,"score_spread":0.3742376563892186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886305294","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07133405,0.0020804866,0.6741716,0.004287989,0.0013086561,0.0013995914,0.08770634,0.11410437,0.043606855],"genre_scores_gemma":[0.24721114,0.0018325928,0.68959177,0.0015070866,0.00047251763,0.0011370837,0.031956308,0.004738599,0.021552874],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996897,0.000103191305,0.00005086736,0.00006310979,0.00006699275,0.000026122665],"domain_scores_gemma":[0.99699926,0.0018702836,0.00024543615,0.00020937304,0.00057295745,0.00010274784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070170517,0.0009860407,0.00029268308,0.0027740821,0.00023869914,0.0015257617,0.00062007795,0.0007087265,0.03609281],"category_scores_gemma":[0.004607853,0.00023965945,0.0004409956,0.0015985189,0.00019738272,0.0010491128,0.0006416928,0.0004942665,0.0047845654],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030423412,0.0004207907,0.01064025,0.0035112174,0.00010444908,0.0028750245,0.004501631,0.009391492,0.07005326,0.0222911,0.36722517,0.50594336],"study_design_scores_gemma":[0.00039756973,0.0006994655,0.038474265,0.0014167589,0.0002091777,0.0033848064,0.00167344,0.14770003,0.08051492,0.01376005,0.71144104,0.00032846973],"about_ca_topic_score_codex":0.0018443387,"about_ca_topic_score_gemma":0.0016779333,"teacher_disagreement_score":0.03609281,"about_ca_system_score_codex":0.00029612757,"about_ca_system_score_gemma":0.00044591815,"threshold_uncertainty_score":0.1207425},"labels":[],"label_agreement":null},{"id":"W2886936867","doi":"10.1109/icosc.2019.8665584","title":"Identifying Protein-Protein Interaction Using Tree LSTM and Structured Attention","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Benchmark (surveying); Computer science; Tree (set theory); Artificial intelligence; Machine learning; Recurrent neural network; Recall; Feature (linguistics); Precision and recall; Deep learning; Pattern recognition (psychology); Artificial neural network; Data mining; Mathematics","score_opus":0.0426747559823257,"score_gpt":0.3231122889954951,"score_spread":0.2804375330131694,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2886936867","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13543917,0.0015390869,0.8517547,0.00057554356,0.00012036858,0.0000763869,0.00061714905,0.00627368,0.0036038447],"genre_scores_gemma":[0.8603149,0.0006651323,0.13327844,0.00031021805,0.00006677359,0.000105565196,0.0012560824,0.00011594235,0.0038869744],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996855,0.000055091103,0.00001911692,0.000119563745,0.000065221466,0.000055626173],"domain_scores_gemma":[0.9994942,0.000239983,0.00008114135,0.00005210946,0.00010436178,0.000028147912],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006073554,0.0008476179,0.0006421198,0.0006408347,0.0002690553,0.0005386465,0.001122035,0.0010656548,0.0013614985],"category_scores_gemma":[0.0016303034,0.00032934448,0.000606329,0.00082957855,0.00031420842,0.0015346243,0.000810811,0.00096448604,0.0008416425],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000453547,0.00042774167,0.003953652,0.00038117636,0.0002984495,0.00053817884,0.00029521788,0.31062853,0.08993337,0.0076794005,0.008011147,0.5773996],"study_design_scores_gemma":[0.0000051613642,0.000045157987,0.00038679584,0.000005929328,0.000019525467,0.000040136314,0.000009548643,0.9908309,0.0047861063,0.0035123997,0.00035188286,0.0000064279316],"about_ca_topic_score_codex":0.005323356,"about_ca_topic_score_gemma":0.007972257,"teacher_disagreement_score":0.005323356,"about_ca_system_score_codex":0.00073684135,"about_ca_system_score_gemma":0.0007352694,"threshold_uncertainty_score":0.010584712},"labels":[],"label_agreement":null},{"id":"W2887611685","doi":"10.1007/978-3-319-99133-7_10","title":"Detecting Low Back Pain from Clinical Narratives Using Machine Learning Approaches","year":2018,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Acronym; Text messaging; Set (abstract data type); Computer science; Artificial intelligence; Subject (documents); Information retrieval; Electronic medical record; Medical record; Plan (archaeology); Natural language processing; Machine learning; Test set; Electronic health record; Medicine; Health care; World Wide Web","score_opus":0.13784280048731015,"score_gpt":0.3456810875785564,"score_spread":0.20783828709124627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2887611685","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.27383748,0.030803306,0.6174573,0.010321466,0.001276608,0.0010199093,0.023021186,0.005713887,0.03654886],"genre_scores_gemma":[0.435849,0.010042322,0.51493806,0.0009712757,0.0006144201,0.00051008526,0.024602177,0.00023856835,0.012234092],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99916613,0.00021201964,0.00012580103,0.000170225,0.00026816007,0.000057769194],"domain_scores_gemma":[0.9961653,0.0030962063,0.00025214098,0.000109068336,0.0003182643,0.000059145692],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012080162,0.00085195474,0.00045604433,0.0034637197,0.00036327742,0.0020890713,0.00082914584,0.0009636823,0.0027027065],"category_scores_gemma":[0.0069936914,0.00025796288,0.0007180801,0.0021830292,0.00034708017,0.0018358425,0.0010234821,0.0010350493,0.0022798153],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013655366,0.00019193365,0.020127304,0.00077827636,0.000112969574,0.0008863795,0.0012056037,0.0045518344,0.013798884,0.003125251,0.018264322,0.9368207],"study_design_scores_gemma":[0.00009024522,0.00052388065,0.092498146,0.00298056,0.0007566459,0.009074596,0.008261139,0.5394847,0.064456046,0.0851569,0.19644043,0.0002767784],"about_ca_topic_score_codex":0.0016716485,"about_ca_topic_score_gemma":0.0032571754,"teacher_disagreement_score":0.0034637197,"about_ca_system_score_codex":0.0005514987,"about_ca_system_score_gemma":0.0008087014,"threshold_uncertainty_score":0.009041488},"labels":[],"label_agreement":null},{"id":"W2890999860","doi":"10.23889/ijpds.v3i4.680","title":"Development of a Concept Dictionary to Standardize Definitions and Classifications While Working With a Common Repository of Linked Administrative Data","year":2018,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute for Clinical Evaluative Sciences","funders":"","keywords":"Computer science; Consistency (knowledge bases); Variety (cybernetics); Quality (philosophy); Data science; Key (lock); Information retrieval; World Wide Web; Artificial intelligence","score_opus":0.27149780613449326,"score_gpt":0.4150943451276689,"score_spread":0.14359653899317565,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2890999860","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0085981395,0.0022806777,0.87522095,0.028444521,0.0070811403,0.013666317,0.010085313,0.0053594657,0.049263533],"genre_scores_gemma":[0.009769039,0.0014626494,0.959085,0.0023116502,0.0006286222,0.0067032184,0.008651469,0.001249966,0.010138378],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.94736296,0.018854275,0.01579411,0.0035816506,0.013171006,0.001236039],"domain_scores_gemma":[0.8401112,0.04100314,0.0159944,0.020338077,0.077278405,0.005274712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06989247,0.0016072959,0.0015910737,0.017848989,0.004400261,0.0152342,0.005631165,0.0031951487,0.014942285],"category_scores_gemma":[0.11135539,0.0015825416,0.0021680468,0.013675444,0.0057216818,0.024297493,0.008692846,0.0072326576,0.0127132535],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014802003,0.0002978506,0.003305187,0.0036023776,0.00005020428,0.0004753008,0.011336852,0.0021072938,0.00417807,0.30276552,0.24231084,0.4294225],"study_design_scores_gemma":[0.000049752038,0.00008565481,0.0012789139,0.003043083,0.000025691175,0.00045030177,0.0027850005,0.0025565927,0.0024032593,0.03701052,0.9501808,0.00013036843],"about_ca_topic_score_codex":0.007047366,"about_ca_topic_score_gemma":0.0054115895,"teacher_disagreement_score":0.06989247,"about_ca_system_score_codex":0.0091643,"about_ca_system_score_gemma":0.034126278,"threshold_uncertainty_score":0.36963117},"labels":[],"label_agreement":null},{"id":"W2891081865","doi":"","title":"Guides: Vancouver Citation Style: Books","year":2016,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Citation; History; Computer science; World Wide Web; Archaeology","score_opus":0.020953829690095956,"score_gpt":0.2791139208443293,"score_spread":0.25816009115423333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891081865","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00029971462,0.0035488463,0.0031915398,0.002513174,0.0029120625,0.00021076459,0.060997795,0.0062220134,0.92010415],"genre_scores_gemma":[0.00070243556,0.0029369395,0.00264547,0.0003416509,0.0002512942,0.000104192775,0.022176605,0.0023167834,0.9685245],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99831116,0.000089146706,0.00012727972,0.00017311494,0.001214012,0.000085206375],"domain_scores_gemma":[0.99025506,0.00097339734,0.00020137089,0.00046620367,0.0072581964,0.0008458055],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001137883,0.0013145405,0.0017374204,0.012802379,0.0033156404,0.010059726,0.002309549,0.001455539,0.60160065],"category_scores_gemma":[0.008892894,0.0011320135,0.000594856,0.027430236,0.0008559614,0.0041761436,0.0016234979,0.0019422567,0.57008857],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000040063937,0.000003440535,0.000028667366,0.00009153527,7.888267e-7,0.000006688818,0.000025642177,0.000025099565,0.000031436455,0.0011662747,0.9786048,0.020011563],"study_design_scores_gemma":[0.000001614676,0.0000013850956,0.00020010251,0.00009323166,0.0000011753046,0.000016052083,0.000035105073,0.000025060437,0.000049135677,0.00048737446,0.999084,0.0000057145257],"about_ca_topic_score_codex":0.14802863,"about_ca_topic_score_gemma":0.35820812,"teacher_disagreement_score":0.9899403,"about_ca_system_score_codex":0.0058957674,"about_ca_system_score_gemma":0.011503799,"threshold_uncertainty_score":0.5682683},"labels":[],"label_agreement":null},{"id":"W2891082052","doi":"10.23889/ijpds.v3i4.782","title":"Using family physician Electronic Medical Record data to measure the pathways of cancer care","year":2018,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Sunnybrook Health Science Centre; Institute for Clinical Evaluative Sciences","funders":"","keywords":"Medicine; Breast cancer; Lung cancer; Specialty; Cancer; Medical record; Internal medicine; Disease; Oncology; Family medicine; Intensive care medicine","score_opus":0.17940420178435806,"score_gpt":0.44479775088113793,"score_spread":0.26539354909677987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891082052","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8529343,0.0014630721,0.0047172057,0.0014887719,0.00006899147,0.0012659152,0.1278463,0.00025548827,0.009959935],"genre_scores_gemma":[0.9206725,0.00094947685,0.013146364,0.0004936252,0.00008861114,0.001386389,0.061130088,0.00003632213,0.0020966607],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9920167,0.002773187,0.0019186474,0.0011109618,0.0017645379,0.00041605203],"domain_scores_gemma":[0.9569471,0.012207253,0.021613054,0.0021920889,0.005932804,0.0011077522],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007813098,0.0003420196,0.00036205383,0.006624371,0.0005154934,0.001444711,0.0006231403,0.00047878007,0.0039762356],"category_scores_gemma":[0.0295028,0.00033026733,0.000552057,0.007925293,0.0002294927,0.0017654216,0.0012011811,0.00051169086,0.00089605566],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013495873,0.00008350341,0.9781924,0.00018630449,0.0001293167,0.00003934777,0.00031959673,0.00029713198,0.00014017917,0.00020024224,0.0033838255,0.016893283],"study_design_scores_gemma":[0.000051275365,0.0002419774,0.98664945,0.00016487282,0.00007613401,0.00015506844,0.0005724183,0.0021096054,0.0005702131,0.00026391266,0.009122413,0.000022649625],"about_ca_topic_score_codex":0.024723483,"about_ca_topic_score_gemma":0.024131201,"teacher_disagreement_score":0.024723483,"about_ca_system_score_codex":0.0015681656,"about_ca_system_score_gemma":0.002110458,"threshold_uncertainty_score":0.04915911},"labels":[],"label_agreement":null},{"id":"W2891264182","doi":"10.23889/ijpds.v3i4.823","title":"A Bayesian Network Model of the Relationships between Chronic Disease Indicators","year":2018,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University Health Centre","funders":"","keywords":"Bayesian network; Context (archaeology); Bayesian probability; Computer science; Causality (physics); Econometrics; Bayes' theorem; Health informatics; Graphical model; Data mining; Machine learning; Data science; Public health; Artificial intelligence; Medicine; Mathematics; Geography","score_opus":0.08500453241448185,"score_gpt":0.3795234618158437,"score_spread":0.2945189294013618,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891264182","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07991031,0.00068941095,0.9021842,0.002527251,0.00009119129,0.00024238978,0.0044278326,0.00042650106,0.009500951],"genre_scores_gemma":[0.7818188,0.0012635711,0.19163135,0.00041506696,0.00022051188,0.0009523533,0.0049193166,0.00011010143,0.018668896],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980116,0.0009794346,0.000069934846,0.00057335274,0.00022653962,0.00013916989],"domain_scores_gemma":[0.9924131,0.005839739,0.00071668706,0.0001774446,0.0006605819,0.00019234353],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004591119,0.00089308433,0.001149789,0.0023660846,0.0006558901,0.001942169,0.0022747936,0.0018943561,0.00748966],"category_scores_gemma":[0.018621823,0.0009252455,0.0010695027,0.002544092,0.001525182,0.0027905963,0.001080521,0.0017529933,0.0008263304],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000104148196,0.000047755693,0.005005088,0.00007530213,0.0000896407,0.0001056412,0.00016435611,0.8791228,0.0002469754,0.09493801,0.0024992994,0.017601034],"study_design_scores_gemma":[0.000022605336,0.000013689011,0.00090984645,0.000020488858,0.000018388879,0.000023488174,0.00001789369,0.9467506,0.00005023466,0.051111925,0.0010468529,0.0000139764925],"about_ca_topic_score_codex":0.043533072,"about_ca_topic_score_gemma":0.03002808,"teacher_disagreement_score":0.043533072,"about_ca_system_score_codex":0.0027638832,"about_ca_system_score_gemma":0.0016831126,"threshold_uncertainty_score":0.086559415},"labels":[],"label_agreement":null},{"id":"W2891337205","doi":"10.23889/ijpds.v3i4.761","title":"Using Biomedical Text as Data and Representation Learning for Identifying Patients with an Osteoarthritis Phenotype in the Electronic Medical Record","year":2018,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Institute for Clinical Evaluative Sciences; University of Toronto","funders":"","keywords":"Artificial intelligence; Machine learning; Computer science; Identification (biology); Supervised learning; Population; Support vector machine; Random forest; Representation (politics); Natural language processing; Medicine; Artificial neural network; Biology","score_opus":0.09758331678051675,"score_gpt":0.44003223897057725,"score_spread":0.3424489221900605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891337205","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.68698937,0.002304911,0.27735126,0.0056956722,0.00031023778,0.00068146404,0.017007157,0.0045833234,0.0050766184],"genre_scores_gemma":[0.7413803,0.0005475289,0.23673004,0.0004364667,0.00021984644,0.00041085135,0.01859206,0.00006264977,0.0016202631],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99689865,0.0017270745,0.00031724904,0.0005796373,0.00034822995,0.00012917368],"domain_scores_gemma":[0.986075,0.010604776,0.0010073264,0.0009924661,0.0011240528,0.00019643868],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035934213,0.0007346793,0.0005031584,0.0031509146,0.00045602757,0.0016649733,0.0010758918,0.0011927203,0.0024337138],"category_scores_gemma":[0.020783179,0.00023466774,0.00058613543,0.0023609418,0.0005629071,0.0021237382,0.0012115701,0.0010190838,0.0012500677],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009279,0.0015585499,0.11011254,0.0008719882,0.0002170039,0.0005497856,0.0007158203,0.060250375,0.011015492,0.0039770296,0.018681478,0.79112196],"study_design_scores_gemma":[0.00015117922,0.00062146835,0.041450754,0.00025853593,0.00015053325,0.00043847322,0.00078201824,0.9062381,0.021632481,0.017254964,0.010918788,0.000102762315],"about_ca_topic_score_codex":0.0032235873,"about_ca_topic_score_gemma":0.0030179997,"teacher_disagreement_score":0.0035934213,"about_ca_system_score_codex":0.00081402913,"about_ca_system_score_gemma":0.000793689,"threshold_uncertainty_score":0.019004047},"labels":[],"label_agreement":null},{"id":"W2891469329","doi":"10.1186/s12911-018-0654-2","title":"Comparison of MetaMap and cTAKES for entity extraction in clinical notes","year":2018,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Unified Medical Language System; Health informatics; Computer science; Information retrieval; Recall; Information extraction; Context (archaeology); Process (computing); Precision and recall; Natural language processing; Data extraction; MEDLINE; Medicine; Pathology; Public health","score_opus":0.1173298955708512,"score_gpt":0.48072275149553845,"score_spread":0.36339285592468723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2891469329","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38726193,0.010330728,0.5016409,0.002772507,0.0008420519,0.0029584675,0.035894968,0.050296996,0.008001461],"genre_scores_gemma":[0.35115832,0.002315604,0.5949071,0.0003918755,0.00015545361,0.0012629644,0.047252826,0.0009037697,0.0016519593],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98169416,0.008094941,0.0036063602,0.0027595283,0.0032096792,0.0006353816],"domain_scores_gemma":[0.9068625,0.071922414,0.0033544216,0.0065733846,0.009858023,0.0014291834],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019201031,0.0025319215,0.0016998539,0.023588002,0.001593787,0.0042304527,0.00197887,0.002405076,0.0023524987],"category_scores_gemma":[0.06715396,0.0008235564,0.0031953051,0.010654662,0.0005972417,0.007301164,0.004455771,0.0014709353,0.001713505],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0077474075,0.0017502786,0.06344677,0.011362128,0.0047977217,0.0015551973,0.0035726496,0.029561318,0.034258395,0.008071179,0.033214908,0.80066204],"study_design_scores_gemma":[0.0009265303,0.0025731921,0.14024316,0.0031319135,0.0039929403,0.005188388,0.006586765,0.5900936,0.12700106,0.02366022,0.09554858,0.001053596],"about_ca_topic_score_codex":0.008370379,"about_ca_topic_score_gemma":0.011311307,"teacher_disagreement_score":0.023588002,"about_ca_system_score_codex":0.001600024,"about_ca_system_score_gemma":0.0035507088,"threshold_uncertainty_score":0.10154593},"labels":[],"label_agreement":null},{"id":"W2892297510","doi":"10.23889/ijpds.v3i4.709","title":"A National Concept Dictionary","year":2018,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Manitoba Health; Institute for Clinical Evaluative Sciences","funders":"","keywords":"Computer science; Flexibility (engineering); Coding (social sciences); Population; World Wide Web; Data science; Information retrieval","score_opus":0.07038370525731832,"score_gpt":0.4236048892272687,"score_spread":0.3532211839699504,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2892297510","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005860864,0.004263919,0.14968364,0.023809448,0.0112745315,0.0055733593,0.08286606,0.006574223,0.710094],"genre_scores_gemma":[0.027040375,0.0069335704,0.46575442,0.010977121,0.0017329687,0.0091181835,0.1759959,0.005015654,0.29743177],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9912396,0.002406657,0.0022495321,0.0011238535,0.0024594825,0.00052092044],"domain_scores_gemma":[0.97387105,0.0040795924,0.001648754,0.003727549,0.014287786,0.0023852456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010590011,0.0010986107,0.0011942749,0.009353544,0.0045523564,0.009583819,0.0036817782,0.0028288784,0.095700495],"category_scores_gemma":[0.02762845,0.0007696057,0.0011036494,0.011583874,0.0030712076,0.015357263,0.0058339564,0.004179564,0.0694042],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009246855,0.0000759973,0.0007127935,0.00095062575,0.000009746468,0.00013676645,0.0019587248,0.00032712665,0.00082212774,0.16121592,0.65056854,0.18312912],"study_design_scores_gemma":[0.000006398012,0.00000986033,0.00019655544,0.0002333546,0.000002105621,0.000105475076,0.0003895645,0.000108062246,0.00015036878,0.0057263994,0.99305856,0.000013277358],"about_ca_topic_score_codex":0.014163362,"about_ca_topic_score_gemma":0.01132567,"teacher_disagreement_score":0.095700495,"about_ca_system_score_codex":0.007911508,"about_ca_system_score_gemma":0.027156709,"threshold_uncertainty_score":0.32015008},"labels":[],"label_agreement":null},{"id":"W2893296467","doi":"10.3233/978-1-61499-910-2-113","title":"The Identity of Dispositions","year":2018,"lang":"en","type":"book-chapter","venue":"Frontiers in artificial intelligence and applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Identity (music); Psychology; Sociology; Philosophy; Aesthetics","score_opus":0.028572315333901283,"score_gpt":0.29791622498490083,"score_spread":0.2693439096509995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2893296467","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14482541,0.0020090973,0.4153298,0.0051348982,0.0013208969,0.0005722029,0.0013006267,0.0012899551,0.42821708],"genre_scores_gemma":[0.87117106,0.00084281387,0.09642197,0.0007868163,0.00028108226,0.000541465,0.0011636855,0.00039079937,0.028400226],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9909978,0.0030110427,0.00093007716,0.002302513,0.001461302,0.0012972063],"domain_scores_gemma":[0.99232566,0.0018641705,0.00084483717,0.002546479,0.0014627884,0.0009560604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059826947,0.0005161328,0.0008071293,0.002374093,0.004867916,0.008057276,0.001612018,0.0024176938,0.00913976],"category_scores_gemma":[0.010659443,0.00064526784,0.0018696381,0.0014609123,0.016708283,0.011387802,0.0075073726,0.0039633242,0.0019334686],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019872836,0.000014167574,0.0010379439,0.000036524456,0.000008783861,0.00009658337,0.0016363648,0.00015368339,0.00027730313,0.9891143,0.00052444555,0.007080049],"study_design_scores_gemma":[0.000029826899,0.000076049226,0.0028257456,0.00019341483,0.00003870422,0.0005022898,0.0029508572,0.0012848137,0.0009781762,0.89303005,0.09803164,0.00005853785],"about_ca_topic_score_codex":0.0030254896,"about_ca_topic_score_gemma":0.0015720468,"teacher_disagreement_score":0.00913976,"about_ca_system_score_codex":0.0034351577,"about_ca_system_score_gemma":0.0033153929,"threshold_uncertainty_score":0.031639934},"labels":[],"label_agreement":null},{"id":"W2895036662","doi":"10.1016/j.jcjd.2018.08.101","title":"Thyroid Nodule Malignancy Rates Within a Health-Care Region with Centralized Pathology","year":2018,"lang":"en","type":"article","venue":"Canadian Journal of Diabetes","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Alberta Bible College","funders":"","keywords":"Medicine; Malignancy; Cytopathology; Thyroid nodules; Nodule (geology); Thyroid; Risk stratification; Radiology; Fine-needle aspiration; Cytology; Pathology; Internal medicine; Biopsy","score_opus":0.011121699648859439,"score_gpt":0.24471443535336454,"score_spread":0.2335927357045051,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2895036662","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9965306,0.0001482555,0.000102115715,0.00015455288,0.0000032734138,0.0000070935303,0.002428537,0.000017752634,0.00060779415],"genre_scores_gemma":[0.9986494,0.00006446712,0.0001394998,0.000026601696,0.000003982987,0.000004185383,0.0009514285,0.0000029650605,0.00015756002],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9991027,0.00015421335,0.00010505933,0.00026991285,0.00017470749,0.00019344],"domain_scores_gemma":[0.9966198,0.00090008834,0.0013122126,0.0001429214,0.0005362089,0.0004887486],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005220302,0.00013770045,0.00029237516,0.001589598,0.0005809499,0.0010215095,0.00057735416,0.00040087124,0.0021071215],"category_scores_gemma":[0.0042172614,0.00021124186,0.0005736196,0.002560038,0.00041170558,0.00055811834,0.00089392124,0.00043494423,0.00025830718],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009777287,0.000019190808,0.9957558,0.000021999494,0.00005998223,0.00015045874,0.00034066028,0.00036559955,0.00022584012,0.0000667528,0.00049444917,0.002401527],"study_design_scores_gemma":[0.0000031995237,0.000021130476,0.99741435,0.000008547471,0.000032685788,0.00019762384,0.0013690473,0.0006704303,0.000070719354,0.000026252648,0.00017951778,0.0000063745283],"about_ca_topic_score_codex":0.43925467,"about_ca_topic_score_gemma":0.45870855,"teacher_disagreement_score":0.43925467,"about_ca_system_score_codex":0.0026042033,"about_ca_system_score_gemma":0.0023632161,"threshold_uncertainty_score":0.87339586},"labels":[],"label_agreement":null},{"id":"W2898281342","doi":"10.1080/17538157.2018.1525734","title":"Bibliometric analysis of the International Medical Informatics Association official journals","year":2018,"lang":"en","type":"article","venue":"Informatics for Health and Social Care","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Health informatics; Bibliometrics; Association (psychology); Data science; Informatics; MEDLINE; Medicine; Political science; Computer science; Public health; Psychology","score_opus":0.025183464552197544,"score_gpt":0.3676348391462326,"score_spread":0.342451374594035,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2898281342","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8215977,0.022633918,0.007411459,0.0033488378,0.0005644959,0.00078629766,0.046492327,0.0011024182,0.096062504],"genre_scores_gemma":[0.96283495,0.008227853,0.0053324522,0.00014145832,0.0007701745,0.0006963447,0.018082567,0.00014846097,0.0037658222],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9698721,0.004202218,0.00584035,0.0015882108,0.017386567,0.0011105611],"domain_scores_gemma":[0.8606041,0.06365544,0.03275988,0.005031333,0.035330452,0.0026188109],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.01310886,0.00066732976,0.0016983099,0.16308972,0.0016946186,0.0073167323,0.0012493047,0.00062232895,0.005437666],"category_scores_gemma":[0.107428625,0.00027216447,0.0013241476,0.21964988,0.001057292,0.0038814978,0.0029955977,0.00053443346,0.0017072399],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043718715,0.00024021203,0.7104043,0.005519858,0.0012406145,0.00037319836,0.0041668997,0.0019112631,0.0020285568,0.009111923,0.021972861,0.24259317],"study_design_scores_gemma":[0.00006848753,0.00019961757,0.89754796,0.0010828387,0.00077027676,0.0012416589,0.0076981196,0.0062295985,0.0028862485,0.0051895995,0.07697669,0.00010887083],"about_ca_topic_score_codex":0.0031877323,"about_ca_topic_score_gemma":0.0023822144,"teacher_disagreement_score":0.83691025,"about_ca_system_score_codex":0.003346159,"about_ca_system_score_gemma":0.00493918,"threshold_uncertainty_score":0.069327116},"labels":[],"label_agreement":null},{"id":"W2899883143","doi":"10.3897/bdj.6.e29616","title":"Incentivising use of structured language in biological descriptions: Author-driven phenotype data and ontology production","year":2018,"lang":"en","type":"article","venue":"Biodiversity Data Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; University of Winnipeg; University of Ottawa; Agriculture and Agri-Food Canada","funders":"","keywords":"Variation (astronomy); Ontology; Computer science; Data science; Scale (ratio); Geography","score_opus":0.17768771626541718,"score_gpt":0.33203328625118544,"score_spread":0.15434556998576826,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899883143","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027216617,0.00038964525,0.918048,0.013503821,0.00041875298,0.0011152745,0.006117308,0.012075663,0.021114875],"genre_scores_gemma":[0.17914282,0.00048187672,0.7885907,0.0014667461,0.00016366193,0.0010364706,0.015698671,0.005532476,0.007886542],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.92435557,0.044702627,0.007869691,0.0065164207,0.0152306,0.001325083],"domain_scores_gemma":[0.6407854,0.19951282,0.0141527625,0.09816593,0.04321893,0.0041641495],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.086153634,0.0011873269,0.0011922795,0.0063280878,0.0020619181,0.0099753495,0.0046136454,0.0028187758,0.006345966],"category_scores_gemma":[0.22078268,0.0017243244,0.002138601,0.00520964,0.005877599,0.02459066,0.015802825,0.004836288,0.0053644255],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076176424,0.0010243357,0.016322564,0.002757207,0.00026429078,0.00135905,0.027604487,0.024353137,0.03155715,0.505449,0.055452317,0.33309472],"study_design_scores_gemma":[0.00019829458,0.00020570256,0.0032311147,0.0009863022,0.00013842428,0.00067490025,0.0053128516,0.12606184,0.043705985,0.414654,0.40442592,0.00040472718],"about_ca_topic_score_codex":0.007047076,"about_ca_topic_score_gemma":0.00773507,"teacher_disagreement_score":0.9138464,"about_ca_system_score_codex":0.0056356685,"about_ca_system_score_gemma":0.00956,"threshold_uncertainty_score":0.4556294},"labels":[],"label_agreement":null},{"id":"W2899915536","doi":"10.1093/nar/gky1032","title":"Human Disease Ontology 2018 update: classification, content and workflow expansion","year":2018,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":557,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"National Cancer Institute; National Human Genome Research Institute","keywords":"Workflow; Disease; Ontology; Biology; Computer science; Data science; Information retrieval; Database; Pathology; Medicine","score_opus":0.12126011907609212,"score_gpt":0.3832616903960068,"score_spread":0.2620015713199147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2899915536","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035706613,0.0086058155,0.2922302,0.08265013,0.017745906,0.004131624,0.2738608,0.18531203,0.09975696],"genre_scores_gemma":[0.030952998,0.008493068,0.33934864,0.010696461,0.0022958943,0.0017687058,0.54873824,0.024565985,0.033139903],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99113095,0.0014794624,0.0019078136,0.0009644123,0.004020585,0.00049684546],"domain_scores_gemma":[0.960052,0.00860422,0.002604093,0.010258328,0.016119411,0.0023619132],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02224844,0.0010188028,0.0009502917,0.008281927,0.0011826014,0.0043329396,0.0031469946,0.0017068548,0.008894024],"category_scores_gemma":[0.057191815,0.0010125547,0.0019094078,0.007107939,0.0007462658,0.008242948,0.006008936,0.003743473,0.010329385],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002164739,0.00018198186,0.0029673164,0.0007943549,0.000056606954,0.00026930848,0.00068425277,0.0010052082,0.0033559247,0.0062149833,0.7603082,0.22394545],"study_design_scores_gemma":[0.000046488294,0.000026136779,0.002007407,0.0003400412,0.000036058813,0.00027832156,0.0001152969,0.002359556,0.0015597729,0.0021371844,0.9910499,0.00004392961],"about_ca_topic_score_codex":0.02697564,"about_ca_topic_score_gemma":0.0344054,"teacher_disagreement_score":0.02697564,"about_ca_system_score_codex":0.0043421015,"about_ca_system_score_gemma":0.0122615695,"threshold_uncertainty_score":0.11766237},"labels":[],"label_agreement":null},{"id":"W2901332105","doi":"10.1093/nar/gky1105","title":"Expansion of the Human Phenotype Ontology (HPO) knowledge base and resources","year":2018,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":737,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; Children's Hospital of Eastern Ontario; University of Toronto; SickKids Foundation; University of Ottawa; Hospital for Sick Children","funders":"NHLBI Division of Intramural Research; National Center for Advancing Translational Sciences; National Human Genome Research Institute; British Heart Foundation; National Cancer Institute; National Institutes of Health; Horizon 2020; National Institute of Allergy and Infectious Diseases; Division of Intramural Research, National Institute of Allergy and Infectious Diseases; European Commission; National Institute for Health and Care Research; E-Rare","keywords":"Biology; Phenotype; Knowledge base; Ontology; Computational biology; Base (topology); Genetics; Gene; World Wide Web; Computer science; Epistemology","score_opus":0.05771103562546586,"score_gpt":0.37128164918889045,"score_spread":0.3135706135634246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2901332105","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028200116,0.006025737,0.6848533,0.023515588,0.0023600839,0.0017016564,0.15342405,0.017139776,0.08277973],"genre_scores_gemma":[0.05707572,0.0060210056,0.6739858,0.0050214874,0.00045951502,0.0012878131,0.245693,0.002065828,0.008389892],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963977,0.0007531208,0.0006420903,0.00065521314,0.0013503678,0.00020145296],"domain_scores_gemma":[0.9890745,0.004302688,0.0007565354,0.0023709764,0.0025493582,0.0009459381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009185796,0.0008339822,0.0009609491,0.0074104774,0.0012923834,0.0030669386,0.0028274271,0.0013373386,0.005690729],"category_scores_gemma":[0.0179107,0.00062878913,0.0017160019,0.007880106,0.0010894388,0.0073913205,0.0053346287,0.003879242,0.0023620327],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035746806,0.0004212268,0.008646867,0.0025366093,0.0002434382,0.0020168403,0.0026087046,0.0060999705,0.007616282,0.16385774,0.27755085,0.528044],"study_design_scores_gemma":[0.000093233335,0.00004545035,0.00602179,0.0013165876,0.00017955092,0.00084758305,0.0005356688,0.011071444,0.0030549474,0.047253508,0.9294767,0.00010359818],"about_ca_topic_score_codex":0.02396446,"about_ca_topic_score_gemma":0.022635335,"teacher_disagreement_score":0.02396446,"about_ca_system_score_codex":0.002379501,"about_ca_system_score_gemma":0.012211675,"threshold_uncertainty_score":0.048579693},"labels":[],"label_agreement":null},{"id":"W2902428389","doi":"10.3897/bdj.6.e29232","title":"Modifier Ontologies for frequency, certainty, degree, and coverage phenotype modifier","year":2018,"lang":"en","type":"article","venue":"Biodiversity Data Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Government of Canada; Agriculture and Agri-Food Canada","funders":"Birmingham Biomedical Research Centre; Imperial Experimental Cancer Medicine Centre; Horizon 2020 Framework Programme; Medical Research Council; Surgical Reconstruction and Microbiology Research Centre; National Institute for Health and Care Research; National Science Foundation","keywords":"Computer science; Ontology; Set (abstract data type); Information retrieval; Class (philosophy); Object (grammar); Interval (graph theory); Degree (music); Certainty; Data mining; Theoretical computer science; Artificial intelligence; Mathematics; Programming language","score_opus":0.12400466249623977,"score_gpt":0.3116825664381551,"score_spread":0.18767790394191536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902428389","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016434127,0.0036381562,0.8468364,0.007972748,0.0006174949,0.0023443399,0.027113728,0.0049902694,0.09005272],"genre_scores_gemma":[0.10009391,0.0035267088,0.83949554,0.0017288843,0.00030945663,0.0037078636,0.03630066,0.0015411017,0.01329587],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9894561,0.002020232,0.002542311,0.0017849605,0.0038862114,0.00031011648],"domain_scores_gemma":[0.9703158,0.01241699,0.003479976,0.005459968,0.0076372167,0.0006900502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013825109,0.0009012875,0.00077837374,0.01329198,0.0020913868,0.0050393785,0.0023337998,0.0016443799,0.010505133],"category_scores_gemma":[0.04097514,0.0007265522,0.0020477644,0.0092546465,0.003300218,0.014774644,0.0037167773,0.0026325313,0.0027698933],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000105909385,0.00008048053,0.0074477275,0.0028767462,0.00009718514,0.00031805725,0.0076475465,0.0025552833,0.005869172,0.7055801,0.04332407,0.2240977],"study_design_scores_gemma":[0.000023733897,0.000028703782,0.0045122756,0.001416418,0.00011385178,0.00054829376,0.0017171586,0.004003175,0.0036599345,0.14072827,0.8431445,0.00010377479],"about_ca_topic_score_codex":0.011339797,"about_ca_topic_score_gemma":0.012245639,"teacher_disagreement_score":0.013825109,"about_ca_system_score_codex":0.0071880314,"about_ca_system_score_gemma":0.0076077776,"threshold_uncertainty_score":0.07311499},"labels":[],"label_agreement":null},{"id":"W2902668122","doi":"10.14569/ijacsa.2018.091101","title":"Exploring Identifiers of Research Articles Related to Food and Disease using Artificial Intelligence","year":2018,"lang":"en","type":"article","venue":"International Journal of Advanced Computer Science and Applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sheridan College; Agriculture and Agri-Food Canada","funders":"","keywords":"Computer science; Identifier; Classifier (UML); Artificial intelligence; Data science; Construct (python library); Machine learning","score_opus":0.164140878476352,"score_gpt":0.4175046342948042,"score_spread":0.2533637558184522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902668122","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73542905,0.05910283,0.06467209,0.0057836315,0.0016926685,0.003647139,0.09353584,0.0020560923,0.03408074],"genre_scores_gemma":[0.5772856,0.025929695,0.3005593,0.0010567743,0.0010262642,0.0023313523,0.08226346,0.00025878532,0.00928884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99432904,0.0010764827,0.0018551908,0.0007202932,0.0017843816,0.00023465893],"domain_scores_gemma":[0.9396372,0.03730092,0.013571442,0.0015268804,0.006909637,0.0010538881],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0055420194,0.00075211225,0.00096775685,0.05527232,0.0012406128,0.0032914462,0.0007780022,0.00079477206,0.0029056822],"category_scores_gemma":[0.03500529,0.00026325288,0.00082996365,0.04537389,0.00064320577,0.0030950839,0.0016237051,0.00047819602,0.0014003508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014250892,0.0004088823,0.2173723,0.021520901,0.000641999,0.008731791,0.0064479676,0.001606018,0.056807917,0.012279901,0.024617616,0.6481396],"study_design_scores_gemma":[0.00025316907,0.0009254608,0.38289103,0.007032854,0.0031793471,0.013483445,0.014252108,0.018812593,0.051806115,0.023912199,0.48309007,0.00036157778],"about_ca_topic_score_codex":0.0019375096,"about_ca_topic_score_gemma":0.0037337774,"teacher_disagreement_score":0.94472766,"about_ca_system_score_codex":0.0014275013,"about_ca_system_score_gemma":0.0040192986,"threshold_uncertainty_score":0.029309392},"labels":[],"label_agreement":null},{"id":"W2902726914","doi":"10.1186/s12911-018-0699-2","title":"Combination of conditional random field with a rule based method in the extraction of PICO elements","year":2018,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":20,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Computer science; Precision and recall; Conditional random field; Information extraction; Element (criminal law); Data mining; Health informatics; Process (computing); Artificial intelligence; Information retrieval; Health care","score_opus":0.024613209495403597,"score_gpt":0.370017778471314,"score_spread":0.34540456897591043,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2902726914","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027638944,0.004510194,0.9569788,0.00069589115,0.00042182664,0.0007080366,0.00159421,0.004976636,0.0024756095],"genre_scores_gemma":[0.18533973,0.0012633714,0.80520934,0.0004829871,0.0003737847,0.0006516372,0.004377433,0.00026709647,0.0020345855],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9935435,0.0023111813,0.0009406841,0.0015102965,0.001502503,0.00019188397],"domain_scores_gemma":[0.96648,0.028297715,0.0010340351,0.001025193,0.0029417193,0.00022127628],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008128014,0.0016180405,0.0018516558,0.00834051,0.00075281353,0.0018321439,0.001775996,0.00206224,0.0025937054],"category_scores_gemma":[0.024908014,0.00044626943,0.0024440645,0.0039284024,0.0007666408,0.0020918725,0.00082265085,0.0015941755,0.0016002686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049215293,0.0003177054,0.006834822,0.0011292017,0.0004604652,0.0008285434,0.00024407265,0.035891116,0.008286726,0.0028302544,0.008015075,0.9346698],"study_design_scores_gemma":[0.00016110335,0.0003244473,0.007577785,0.0004860703,0.00058677525,0.001369769,0.00016772158,0.94208103,0.019434793,0.01428762,0.013350977,0.00017176975],"about_ca_topic_score_codex":0.0066043218,"about_ca_topic_score_gemma":0.0061711064,"teacher_disagreement_score":0.00834051,"about_ca_system_score_codex":0.0008690266,"about_ca_system_score_gemma":0.00216945,"threshold_uncertainty_score":0.04298556},"labels":[],"label_agreement":null},{"id":"W2903220936","doi":"","title":"Research guides: Grey Literature and Statistics for Dentistry: Home","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Grey literature; Oral health; Geography; Statistics; Dentistry; MEDLINE; Medicine; Political science; Mathematics","score_opus":0.07723222167012465,"score_gpt":0.3789187941579624,"score_spread":0.30168657248783776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2903220936","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021280523,0.010642789,0.11211495,0.019746793,0.0023459229,0.005461326,0.49726298,0.04920955,0.30108756],"genre_scores_gemma":[0.007201315,0.02203649,0.50715834,0.0065738205,0.0013530483,0.009404014,0.23755008,0.019393014,0.18932989],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99607855,0.0010596331,0.0010783778,0.00022859244,0.0014324137,0.00012240668],"domain_scores_gemma":[0.93034166,0.048775278,0.0032354381,0.0022735144,0.013997222,0.0013768274],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0071407673,0.0012075319,0.0018774254,0.014571399,0.0010624516,0.00574292,0.002534195,0.0022405118,0.31039512],"category_scores_gemma":[0.05801104,0.001594899,0.0011519147,0.022638872,0.00076256675,0.0061663585,0.0033248684,0.0021430883,0.19461152],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003565188,0.000034551296,0.00011914601,0.0028470715,0.000011423577,0.000073323616,0.00044433723,0.00020243098,0.00017256071,0.005138944,0.8836846,0.10723604],"study_design_scores_gemma":[0.00004331386,0.0000126685745,0.00037322615,0.0018526142,0.000009809591,0.00009538583,0.00028258652,0.0001916421,0.00016643328,0.006384598,0.9905618,0.000025998097],"about_ca_topic_score_codex":0.008657836,"about_ca_topic_score_gemma":0.016782861,"teacher_disagreement_score":0.31039512,"about_ca_system_score_codex":0.0020653768,"about_ca_system_score_gemma":0.009721655,"threshold_uncertainty_score":0.9836377},"labels":[],"label_agreement":null},{"id":"W2903411464","doi":"10.4018/ijeach.2019010108","title":"Cascading Workflow of Healthcare Services","year":2018,"lang":"en","type":"article","venue":"International Journal of Extreme Automation and Connectivity in Healthcare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thunder Bay Regional Health Sciences Centre; Lakehead University","funders":"","keywords":"Workflow; Interoperability; Computer science; JSON; Health care; Scalability; Knowledge management; Data science; Software engineering; Process management; World Wide Web; Database; Business","score_opus":0.03015221905081906,"score_gpt":0.3298946715821289,"score_spread":0.2997424525313098,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2903411464","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07664002,0.0011141811,0.85838145,0.0032438918,0.0003929999,0.0024693175,0.0026857564,0.020972552,0.034099877],"genre_scores_gemma":[0.41968346,0.001133708,0.55306935,0.000536263,0.00011376819,0.00068729796,0.005545259,0.0013211262,0.017909732],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99142283,0.0018789734,0.0010820156,0.0019916827,0.0029205466,0.0007040374],"domain_scores_gemma":[0.991808,0.0022940075,0.00042685305,0.0025149074,0.002112841,0.00084349496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007812199,0.001567561,0.0010749532,0.0046791276,0.0029359176,0.006371523,0.002304442,0.0014979023,0.009462998],"category_scores_gemma":[0.01510595,0.001002777,0.002204449,0.0045097307,0.0013038601,0.00377357,0.0071155042,0.0020208326,0.0042132847],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009334496,0.00069801963,0.031603005,0.0015987841,0.000822796,0.005741854,0.009549577,0.13179064,0.028772738,0.14795037,0.04802391,0.59251493],"study_design_scores_gemma":[0.00017279228,0.00024841286,0.009544395,0.00048693008,0.00054191926,0.0011967209,0.0030941626,0.49326584,0.023045534,0.2708764,0.19719423,0.00033268193],"about_ca_topic_score_codex":0.02363064,"about_ca_topic_score_gemma":0.01658903,"teacher_disagreement_score":0.02363064,"about_ca_system_score_codex":0.002299012,"about_ca_system_score_gemma":0.006792264,"threshold_uncertainty_score":0.046986222},"labels":[],"label_agreement":null},{"id":"W2904022159","doi":"10.1101/500686","title":"Text-mining clinically relevant cancer biomarkers for curation into the CIViC database","year":2018,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"National Cancer Institute; National Human Genome Research Institute; National Institutes of Health","keywords":"Data curation; Computer science; Precision medicine; Construct (python library); Cancer; Resource (disambiguation); MEDLINE; Information retrieval; Data science; Medicine; Pathology; Biology; Internal medicine","score_opus":0.024914695094531192,"score_gpt":0.2972968484878694,"score_spread":0.2723821533933382,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2904022159","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015040894,0.007543653,0.04089352,0.0027540505,0.00047215255,0.0021782543,0.9037026,0.018147882,0.009267001],"genre_scores_gemma":[0.022375362,0.0030259872,0.15413244,0.0011519747,0.00019671497,0.0020883742,0.81440014,0.00095854903,0.0016704563],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9958912,0.00072361156,0.0013685767,0.0009986868,0.00085739157,0.00016048505],"domain_scores_gemma":[0.9728276,0.014343959,0.0041311686,0.0024250855,0.005336606,0.00093547825],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068569356,0.0021899554,0.0025155724,0.027861059,0.001647747,0.0036860537,0.0029927306,0.0030002992,0.013652777],"category_scores_gemma":[0.032867696,0.0009960444,0.0019159156,0.015812699,0.0006748128,0.0032761716,0.0032858066,0.002243345,0.0090236],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012584843,0.00061732344,0.017549302,0.049958386,0.001048923,0.006217251,0.0025446727,0.005400245,0.034582976,0.01237898,0.6263517,0.24209175],"study_design_scores_gemma":[0.0004944646,0.00038643132,0.023686528,0.006927995,0.0016783153,0.002615572,0.0015681724,0.019375253,0.030552117,0.013153057,0.89926726,0.00029488278],"about_ca_topic_score_codex":0.005464245,"about_ca_topic_score_gemma":0.010825212,"teacher_disagreement_score":0.027861059,"about_ca_system_score_codex":0.0025147805,"about_ca_system_score_gemma":0.008056015,"threshold_uncertainty_score":0.045673072},"labels":[],"label_agreement":null},{"id":"W2905553692","doi":"10.1186/s13321-018-0319-2","title":"SIA: a scalable interoperable annotation server for biomedical named entities","year":2018,"lang":"en","type":"article","venue":"Journal of Cheminformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Wirtschaft und Energie; Banting and Best Diabetes Centre, University of Toronto; Bundesministerium für Bildung und Forschung","keywords":"Annotation; Interoperability; Workflow; Scalability; Server; Information extraction; Named-entity recognition","score_opus":0.013462996530450582,"score_gpt":0.27915431022398707,"score_spread":0.2656913136935365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2905553692","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003671203,0.00077404076,0.3540823,0.0015075714,0.00047560327,0.0007816495,0.07350397,0.55823874,0.0069649257],"genre_scores_gemma":[0.03189932,0.0010794302,0.35832286,0.0013463984,0.00030271176,0.001666068,0.55846334,0.035667393,0.011252467],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953324,0.00078091864,0.00079041324,0.0010293645,0.0017321702,0.00033461436],"domain_scores_gemma":[0.99025327,0.0022768828,0.001003957,0.003070474,0.0020371107,0.001358263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0084144175,0.0032732126,0.0023334324,0.008837342,0.0022746224,0.005433147,0.005276086,0.0027500181,0.019135471],"category_scores_gemma":[0.017397791,0.0018836188,0.0026274784,0.008321518,0.0012548723,0.008022275,0.008962981,0.004550034,0.04316473],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001973238,0.00054495584,0.005692155,0.0028678677,0.00074857683,0.0012118061,0.0011190607,0.004570455,0.047524355,0.02714977,0.7294292,0.17716864],"study_design_scores_gemma":[0.0006183413,0.00023789502,0.0076083546,0.00064868235,0.0003367625,0.001474851,0.00061223086,0.11746537,0.046129193,0.07848434,0.7457469,0.00063707307],"about_ca_topic_score_codex":0.0038503474,"about_ca_topic_score_gemma":0.0030804232,"teacher_disagreement_score":0.019135471,"about_ca_system_score_codex":0.0016750618,"about_ca_system_score_gemma":0.0058957757,"threshold_uncertainty_score":0.064014554},"labels":[],"label_agreement":null},{"id":"W2906181611","doi":"10.3389/fninf.2018.00085","title":"National Neuroinformatics Framework for Canadian Consortium on Neurodegeneration in Aging (CCNA)","year":2018,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University; Jewish General Hospital; Institut Universitaire de Gériatrie de Montréal; Montreal Neurological Institute and Hospital","funders":"Consortium canadien en neurodégénérescence associée au vieillissement","keywords":"Neuroinformatics; Computer science; World Wide Web; Medicine; Data science","score_opus":0.019795526947114134,"score_gpt":0.2795183087001601,"score_spread":0.259722781753046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2906181611","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005833545,0.0063314587,0.32004607,0.058100626,0.0025370482,0.004101983,0.20484035,0.0508641,0.34734488],"genre_scores_gemma":[0.043180387,0.0065575833,0.49894658,0.01185194,0.000676505,0.003207624,0.34162167,0.008566713,0.08539102],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9762056,0.0043119816,0.0016862323,0.0028892169,0.011679495,0.0032274867],"domain_scores_gemma":[0.8910357,0.0071010212,0.0017227365,0.009505376,0.07998293,0.010652244],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035171617,0.0020675096,0.001553852,0.009410034,0.007642763,0.015982008,0.0112948725,0.0034255695,0.025642822],"category_scores_gemma":[0.051704925,0.0012025869,0.0018750072,0.012064152,0.00373299,0.0067933127,0.01331885,0.004083925,0.018447671],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027296372,0.000110401605,0.0034297062,0.0005285127,0.00011210125,0.00024658337,0.001816729,0.003028869,0.0015978229,0.13594697,0.73123854,0.121670805],"study_design_scores_gemma":[0.00004541993,0.000021019725,0.0025405409,0.0004763058,0.000042498283,0.00009083722,0.0007085797,0.0048862025,0.001196014,0.021229636,0.9686593,0.00010370879],"about_ca_topic_score_codex":0.92533857,"about_ca_topic_score_gemma":0.9000664,"teacher_disagreement_score":0.93513626,"about_ca_system_score_codex":0.064863764,"about_ca_system_score_gemma":0.2821337,"threshold_uncertainty_score":0.47062176},"labels":[],"label_agreement":null},{"id":"W2911446348","doi":"10.1186/s13023-018-0980-6","title":"An ontological foundation for ocular phenotypes and rare eye diseases","year":2019,"lang":"en","type":"letter","venue":"Orphanet Journal of Rare Diseases","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; Všeobecná Fakultní Nemocnice v Praze; Moorfields Eye Hospital NHS Foundation Trust; Uniwersytet Medyczny w Lublinie; Université de Lausanne; Hospital for Sick Children; Universiteit Gent; Newcastle University; Universiteit van Amsterdam; Institut National de la Santé et de la Recherche Médicale; Assistance publique-Hôpitaux de Paris; Leids Universitair Medisch Centrum; Universiteit Leiden; National Institute for Health and Care Research; Universitair Ziekenhuis Gent; Université de Montpellier; Ospedale Pediatrico Bambino Gesù; Hadassah Medical Organization; Universität Basel; European Commission; Eberhard Karls Universität Tübingen; Johns Hopkins University; Rigshospitalet; Cleveland Clinic; Johannes Gutenberg-Universität Mainz; Universidade de Coimbra; Univerzita Karlova v Praze","keywords":"Foundation (evidence); Phenotype; Human genetics; Genetics; Medicine; Optometry; Biology; Political science; Law; Gene","score_opus":0.014809440085015795,"score_gpt":0.28429999119922117,"score_spread":0.2694905511142054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911446348","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033114407,0.0036566374,0.85183156,0.013668788,0.0008592049,0.0012727336,0.015863152,0.0014697904,0.07826372],"genre_scores_gemma":[0.24774222,0.0037920491,0.7089929,0.0024177562,0.0004886061,0.0014257873,0.027591394,0.00028952886,0.007259798],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99582785,0.001475657,0.0007588034,0.0005814251,0.0010643172,0.00029185688],"domain_scores_gemma":[0.9918007,0.0040415223,0.0010007159,0.0010819564,0.0016016332,0.0004733809],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060063535,0.00056337455,0.0005074063,0.006982137,0.002665045,0.004345325,0.0013009532,0.0015238181,0.0031639675],"category_scores_gemma":[0.011696063,0.00043585425,0.0018290821,0.0055565448,0.0026430779,0.006380489,0.004631509,0.00218979,0.0008751024],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000731597,0.000108209264,0.007190313,0.00072439853,0.00010327242,0.001184141,0.0062796124,0.0040006745,0.0030508963,0.8700003,0.018078294,0.08920676],"study_design_scores_gemma":[0.000046780984,0.00004310159,0.010628642,0.0015045678,0.00022023651,0.0017651145,0.005020031,0.021615831,0.0017059966,0.5165892,0.44077125,0.00008925655],"about_ca_topic_score_codex":0.017487407,"about_ca_topic_score_gemma":0.017460333,"teacher_disagreement_score":0.017487407,"about_ca_system_score_codex":0.0029940086,"about_ca_system_score_gemma":0.0077673243,"threshold_uncertainty_score":0.034771264},"labels":[],"label_agreement":null},{"id":"W2911525976","doi":"10.1109/bibm.2018.8621195","title":"Boundary Detection by Determining the Difference of Classification Probabilities of Sequences: Topic Segmentation of Clinical Notes","year":2018,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Support vector machine; Artificial intelligence; Segmentation; Naive Bayes classifier; Computer science; Boundary (topology); Machine learning; Bayes' theorem; Sequence (biology); Natural language processing; Pattern recognition (psychology); Mathematics; Bayesian probability","score_opus":0.07109220662878026,"score_gpt":0.3652287898102724,"score_spread":0.29413658318149216,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2911525976","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33760318,0.001612419,0.65300685,0.0006342218,0.00023622465,0.0004510751,0.00076771376,0.00317872,0.002509541],"genre_scores_gemma":[0.734486,0.0003103073,0.25978196,0.00020179368,0.00022861683,0.00027401428,0.002776875,0.00022360566,0.0017167992],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9966928,0.00097341114,0.00030753185,0.0010143148,0.0006896882,0.00032223252],"domain_scores_gemma":[0.99121726,0.00551978,0.00082914357,0.0005742151,0.0014694728,0.00039011618],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044120145,0.00097178534,0.0014110879,0.0040515433,0.00078904827,0.0016244961,0.001004004,0.0017665396,0.0011338358],"category_scores_gemma":[0.013911672,0.00036259147,0.0011416613,0.0020106407,0.00074782915,0.002469123,0.0012147797,0.00154736,0.0011206532],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023852969,0.00064791343,0.06000417,0.00039007197,0.0002699437,0.00036476823,0.0013782532,0.03740149,0.061086837,0.0033271702,0.008132651,0.8246115],"study_design_scores_gemma":[0.00014178576,0.00076391053,0.043128874,0.00007965955,0.00014863214,0.00067191315,0.00070758903,0.9010442,0.037741836,0.009837327,0.005622663,0.00011150542],"about_ca_topic_score_codex":0.0032716542,"about_ca_topic_score_gemma":0.0022391777,"teacher_disagreement_score":0.0044120145,"about_ca_system_score_codex":0.0006759089,"about_ca_system_score_gemma":0.0011996815,"threshold_uncertainty_score":0.023333192},"labels":[],"label_agreement":null},{"id":"W2912907847","doi":"10.1016/j.jbi.2019.103114","title":"Automatic ICD code assignment of Chinese clinical notes based on multilayer attention BiRNN","year":2019,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":66,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Higher Education Discipline Innovation Project; National Natural Science Foundation of China","keywords":"Computer science; Code (set theory); Feature engineering; Word embedding; Feature (linguistics); Artificial intelligence; Natural language processing; Block (permutation group theory); Representation (politics); Artificial neural network; Semantics (computer science); Word (group theory); Recurrent neural network; Hamming code; Embedding; Deep learning; Block code; Linguistics; Programming language; Algorithm; Decoding methods","score_opus":0.02064373706993933,"score_gpt":0.33888980374430683,"score_spread":0.3182460666743675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912907847","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7094874,0.0048764064,0.24733865,0.002256937,0.001159214,0.00081390486,0.01805188,0.0073554795,0.008660165],"genre_scores_gemma":[0.84662586,0.0011270104,0.1209034,0.0003374996,0.00028551064,0.00033060266,0.026284877,0.0001260712,0.003979124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99877733,0.00015874508,0.0002278451,0.00037347246,0.00027178723,0.00019075754],"domain_scores_gemma":[0.99768937,0.0007535522,0.00021565073,0.00018717664,0.0010219328,0.00013244312],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010141226,0.0005741303,0.0006166988,0.005476867,0.00096269354,0.0008770082,0.000793725,0.0006075577,0.0020292655],"category_scores_gemma":[0.0037114539,0.00018806173,0.0006142591,0.002491791,0.00025508448,0.0008955427,0.0011034381,0.0006961726,0.0010130545],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009337997,0.00054398674,0.11564579,0.0007332645,0.00016077417,0.0012194081,0.0006723348,0.0064526317,0.0319125,0.002164459,0.033130348,0.8064307],"study_design_scores_gemma":[0.00017893914,0.0003749518,0.16417472,0.00033888,0.0008496188,0.002201351,0.0016871024,0.76071197,0.036533613,0.010251292,0.022514056,0.00018345144],"about_ca_topic_score_codex":0.02253406,"about_ca_topic_score_gemma":0.027814617,"teacher_disagreement_score":0.02253406,"about_ca_system_score_codex":0.0009972298,"about_ca_system_score_gemma":0.0021591615,"threshold_uncertainty_score":0.044805825},"labels":[],"label_agreement":null},{"id":"W2912922979","doi":"10.1101/532523","title":"Gene Info: Easy retrieval of gene product information on any website","year":2019,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Sinai Health System; Lunenfeld-Tanenbaum Research Institute","funders":"","keywords":"Product (mathematics); Quality (philosophy); Reading (process); Scale (ratio); Genomics","score_opus":0.011577254811598383,"score_gpt":0.22399718737558014,"score_spread":0.21241993256398176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2912922979","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053406116,0.0010109443,0.050792716,0.0016208746,0.0005835583,0.0005406068,0.45338506,0.44419155,0.042534184],"genre_scores_gemma":[0.029478565,0.0016344124,0.08257817,0.0018772363,0.0005249212,0.001050715,0.74556863,0.07451156,0.062775776],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99916553,0.0000632836,0.00007451797,0.00016512175,0.00040278002,0.00012866601],"domain_scores_gemma":[0.99861646,0.00026821994,0.00011549712,0.0004965885,0.0002496956,0.00025344858],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00092763285,0.0026444294,0.0014399079,0.0040360666,0.00070282485,0.0024544543,0.0019234052,0.0020476007,0.21625113],"category_scores_gemma":[0.003173615,0.0011830826,0.001040778,0.0032922754,0.00048808576,0.002876021,0.0038134411,0.0017767779,0.30064154],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004931474,0.00010985709,0.0010060946,0.00081395486,0.000069322545,0.00027729364,0.00010372166,0.00018806067,0.016419074,0.0022437405,0.9178774,0.060398344],"study_design_scores_gemma":[0.000408707,0.00011796034,0.00930456,0.00026648567,0.000060550425,0.0008945049,0.00015291532,0.0029103437,0.033946704,0.011308766,0.9404675,0.0001610647],"about_ca_topic_score_codex":0.0017342906,"about_ca_topic_score_gemma":0.0031995072,"teacher_disagreement_score":0.21625113,"about_ca_system_score_codex":0.0008598874,"about_ca_system_score_gemma":0.0009487564,"threshold_uncertainty_score":0.7234321},"labels":[],"label_agreement":null},{"id":"W2913179481","doi":"","title":"A (acronyms)","year":2004,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Acronym; Computer science; Natural language processing; Artificial intelligence; Set (abstract data type); Phrase; Lexicon; TRACE (psycholinguistics); Linguistics; Programming language","score_opus":0.0094971380826072,"score_gpt":0.2572173758345895,"score_spread":0.2477202377519823,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913179481","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013424051,0.009216023,0.114767216,0.0045055593,0.013485493,0.0008672808,0.039597332,0.013372742,0.7907644],"genre_scores_gemma":[0.1683037,0.009041406,0.21480426,0.0051647383,0.0036700505,0.0011704357,0.052015256,0.0052448483,0.5405853],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973022,0.0006484628,0.00042915234,0.00068944256,0.0007225217,0.0002081835],"domain_scores_gemma":[0.99681133,0.0005764459,0.00038145375,0.00081849314,0.0012129188,0.00019944427],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010611463,0.0014909067,0.00074897916,0.0038563854,0.001992797,0.0050995788,0.0014219974,0.001326556,0.1624771],"category_scores_gemma":[0.004964591,0.00038897147,0.00092885736,0.004633609,0.0014111963,0.0075838924,0.0029245063,0.001899889,0.12809914],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003008167,0.00008570606,0.00242871,0.0013114489,0.0000864061,0.0007938516,0.0010201233,0.00066770194,0.0069440547,0.20477521,0.42505342,0.35653242],"study_design_scores_gemma":[0.000007040005,0.000025054333,0.00042099477,0.00009066029,0.000009369935,0.00045989774,0.00023012588,0.00028848485,0.0007775402,0.0116245765,0.9860436,0.00002260698],"about_ca_topic_score_codex":0.0018962792,"about_ca_topic_score_gemma":0.001962348,"teacher_disagreement_score":0.83752286,"about_ca_system_score_codex":0.0010609274,"about_ca_system_score_gemma":0.0011864393,"threshold_uncertainty_score":0.54354006},"labels":[],"label_agreement":null},{"id":"W2913686418","doi":"10.3166/isi.23.2.39-59","title":"Moteur de révision d’ontologie en SHIQ","year":2018,"lang":"fr","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.015222628499117407,"score_gpt":0.2696375714238626,"score_spread":0.25441494292474515,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913686418","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0414072,0.0044905124,0.75594723,0.027258467,0.013185328,0.0014004823,0.015931306,0.022444239,0.11793525],"genre_scores_gemma":[0.2273408,0.0071376427,0.5236589,0.0044875373,0.005390024,0.0014955349,0.036630373,0.011697639,0.18216161],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9888461,0.0037024708,0.0011749318,0.0015741158,0.0042493735,0.000453032],"domain_scores_gemma":[0.98229206,0.0043309457,0.00060180476,0.0030344003,0.009160006,0.00058078417],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010300471,0.0014902509,0.0011446134,0.0057642707,0.0030010112,0.007873875,0.0015136574,0.0019034831,0.032722797],"category_scores_gemma":[0.026961448,0.00079400674,0.002447144,0.004076343,0.0029443677,0.009898167,0.004237364,0.0050334237,0.015439831],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006553356,0.00026002363,0.0073694405,0.0013941668,0.0002107046,0.0017063543,0.005503388,0.003959908,0.01565647,0.24791242,0.21358682,0.50178504],"study_design_scores_gemma":[0.00008266804,0.00009029519,0.0045186426,0.00045987457,0.00008162439,0.0014498712,0.0020315214,0.011070822,0.011619203,0.05140792,0.9170315,0.0001560536],"about_ca_topic_score_codex":0.029988892,"about_ca_topic_score_gemma":0.011737524,"teacher_disagreement_score":0.032722797,"about_ca_system_score_codex":0.0033411568,"about_ca_system_score_gemma":0.007365432,"threshold_uncertainty_score":0.1094687},"labels":[],"label_agreement":null},{"id":"W2913815335","doi":"10.13020/d6bm3v","title":"integrated DIetary Supplement Knowledge base (iDISK)","year":2017,"lang":"en","type":"dataset","venue":"University of Minnesota Digital Conservancy (University of Minnesota)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Documentation; Computer science; Software; File format; World Wide Web; Information retrieval; Software engineering; Database; Operating system","score_opus":0.024647120189146072,"score_gpt":0.23844172739305286,"score_spread":0.21379460720390678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2913815335","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012210369,0.0005835347,0.00096638553,0.00020142902,0.000027379705,0.000093277486,0.9933923,0.00089976325,0.002614866],"genre_scores_gemma":[0.002060226,0.0004916759,0.0032234392,0.000118805285,0.0000054250118,0.0002032225,0.9931176,0.00005674527,0.0007229394],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992766,0.00011428169,0.00012930857,0.00022206227,0.0001982929,0.000059398244],"domain_scores_gemma":[0.9987123,0.00045927175,0.00017393695,0.00024526188,0.00028827868,0.0001210317],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010069823,0.0014418234,0.0010798841,0.0053319274,0.0005638836,0.0016959438,0.0028285861,0.0017036565,0.019899221],"category_scores_gemma":[0.004952094,0.0006039991,0.0012477451,0.007973849,0.00034349464,0.0016191439,0.00190544,0.0015638779,0.013951382],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063036877,0.00028638492,0.006624759,0.0073343283,0.00032643022,0.00042423193,0.00016168455,0.0043021557,0.001426528,0.0049261726,0.92201746,0.051539425],"study_design_scores_gemma":[0.00026254935,0.000043470733,0.0047710855,0.0006648298,0.00016619358,0.00024704015,0.00010560714,0.0017881106,0.0012796166,0.0028235172,0.9878051,0.00004278164],"about_ca_topic_score_codex":0.02161202,"about_ca_topic_score_gemma":0.038394872,"teacher_disagreement_score":0.02161202,"about_ca_system_score_codex":0.0017343169,"about_ca_system_score_gemma":0.0038196524,"threshold_uncertainty_score":0.06656951},"labels":[],"label_agreement":null},{"id":"W2914171828","doi":"10.1093/database/bay147","title":"Overview of the BioCreative VI Precision Medicine Track: mining protein interactions and mutations for precision medicine","year":2018,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":75,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Institute for Research in Immunology and Cancer","funders":"U.S. National Library of Medicine; National Institute of General Medical Sciences; National Cancer Institute; National Institutes of Health","keywords":"Computer science; Task (project management); Precision medicine; Triage; F1 score; Relationship extraction; Precision and recall; Annotation; Information extraction; Named-entity recognition; Information retrieval; Relation (database); Personalized medicine; Natural language processing; Data science; Data mining; Artificial intelligence; Bioinformatics; Medicine; Genetics; Biology","score_opus":0.07326093043933397,"score_gpt":0.3871233921750868,"score_spread":0.31386246173575283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914171828","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035572726,0.030569758,0.30553755,0.006009399,0.0013655289,0.0048539913,0.4222121,0.14812417,0.0457548],"genre_scores_gemma":[0.022887422,0.005556324,0.3137175,0.001317606,0.0004376132,0.002119886,0.63377506,0.003149193,0.017039401],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946636,0.00071374705,0.0006103887,0.0014750483,0.0021694256,0.00036772116],"domain_scores_gemma":[0.9913424,0.0026621327,0.000747712,0.0014191949,0.0029470634,0.0008814612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076008355,0.0022829596,0.0016207603,0.011905538,0.0014976878,0.005045407,0.0032464694,0.0021276716,0.015125301],"category_scores_gemma":[0.0103117395,0.0012639721,0.0024261915,0.008776105,0.00040506414,0.0045349346,0.0024544518,0.002001177,0.016549544],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010111572,0.00067164906,0.015609771,0.003946405,0.0006113498,0.00033505083,0.00049329665,0.006309255,0.019873122,0.00404181,0.5350568,0.41204032],"study_design_scores_gemma":[0.00050471874,0.0012565354,0.027476663,0.0009612133,0.0005565912,0.0011741185,0.0002564573,0.061174177,0.025231758,0.0062370882,0.87490696,0.00026369301],"about_ca_topic_score_codex":0.016743204,"about_ca_topic_score_gemma":0.022461848,"teacher_disagreement_score":0.016743204,"about_ca_system_score_codex":0.0019821443,"about_ca_system_score_gemma":0.0047437227,"threshold_uncertainty_score":0.050599158},"labels":[],"label_agreement":null},{"id":"W2914441380","doi":"10.3390/ijns5010009","title":"Building a Newborn Screening Information Management System from Theory to Practice","year":2019,"lang":"en","type":"article","venue":"International Journal of Neonatal Screening","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Ottawa; Newborn Screening Ontario; Children's Hospital of Eastern Ontario","funders":"","keywords":"Flexibility (engineering); Procurement; Process management; Process (computing); Knowledge management; Information system; Work (physics); Engineering management; Computer science; Business; Engineering; Marketing","score_opus":0.007855951743561695,"score_gpt":0.2779725315462516,"score_spread":0.27011657980268994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2914441380","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026987897,0.0008574031,0.90992206,0.0074660853,0.0003236887,0.0018250021,0.00091443106,0.03514424,0.016559297],"genre_scores_gemma":[0.08137781,0.0006488242,0.9099085,0.00094373483,0.00007742792,0.00054326863,0.0021212995,0.0006235666,0.0037556258],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99262035,0.0019942727,0.0009878977,0.0014815386,0.0024734947,0.00044242237],"domain_scores_gemma":[0.9879872,0.002842861,0.000702172,0.002331713,0.0043500564,0.0017859305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015616217,0.00057691743,0.0007780849,0.0030868382,0.0016855869,0.006263073,0.003243616,0.002137597,0.004960035],"category_scores_gemma":[0.018227817,0.0009240003,0.000842316,0.0019087791,0.0013123882,0.008747224,0.00447255,0.0028795335,0.004048183],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003828908,0.001274437,0.020330772,0.0013548572,0.00018439707,0.0020661724,0.006300374,0.014972479,0.025083596,0.054830108,0.057396226,0.81582373],"study_design_scores_gemma":[0.00023830369,0.0012193226,0.016100897,0.0034674155,0.0002921262,0.0027264364,0.0047499337,0.2640427,0.041417178,0.072818704,0.59236604,0.0005609114],"about_ca_topic_score_codex":0.0063207103,"about_ca_topic_score_gemma":0.0039525433,"teacher_disagreement_score":0.015616217,"about_ca_system_score_codex":0.0037081372,"about_ca_system_score_gemma":0.0091917785,"threshold_uncertainty_score":0.08258748},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"not_applicable","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"methods","about_ca_system":false,"about_ca_topic":false,"confidence":"high"}],"label_agreement":"split"},{"id":"W2915926009","doi":"10.1186/2041-1480-4-s1-i1","title":"Selected papers from the 15th Annual Bio-Ontologies Special Interest Group Meeting","year":2013,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Biomedicine; Ontology; Computer science; Presentation (obstetrics); Data science; World Wide Web; Semantic Web; Annotation; Open Biomedical Ontologies; Information retrieval; Ontology-based data integration; Bioinformatics; Ontology alignment; Artificial intelligence; Medicine","score_opus":0.0137067060444557,"score_gpt":0.2471825314170106,"score_spread":0.2334758253725549,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2915926009","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001463374,0.038418055,0.0027775127,0.16426931,0.7436754,0.001039406,0.0038673978,0.00029943002,0.044190127],"genre_scores_gemma":[0.0074722273,0.052834623,0.0038112479,0.06121316,0.50583667,0.001492909,0.008859452,0.00073593133,0.35774365],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953703,0.0006521696,0.00042602766,0.000440331,0.002639887,0.0004712729],"domain_scores_gemma":[0.9831149,0.001858782,0.00086203316,0.0003327941,0.009864506,0.003966994],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009320781,0.002460534,0.002101571,0.0075778584,0.0038647659,0.008594309,0.0023685428,0.0051282495,0.08863904],"category_scores_gemma":[0.0134887425,0.0006876752,0.0019368413,0.0053224093,0.0008558529,0.00430669,0.004971424,0.0040247045,0.03737246],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006173955,0.000023878629,0.00009873952,0.0001903265,0.000007987416,0.000089047666,0.000043652817,0.000021178366,0.00023626728,0.00019477658,0.9829495,0.016082909],"study_design_scores_gemma":[0.000020203113,0.000025784566,0.00058145134,0.0002822763,0.000015660666,0.00004702779,0.00014093256,0.00005184693,0.0001207492,0.0004354616,0.9982663,0.000012387779],"about_ca_topic_score_codex":0.00423294,"about_ca_topic_score_gemma":0.015268232,"teacher_disagreement_score":0.08863904,"about_ca_system_score_codex":0.0048202653,"about_ca_system_score_gemma":0.0059582186,"threshold_uncertainty_score":0.29652715},"labels":[],"label_agreement":null},{"id":"W2916913049","doi":"10.3389/fphys.2019.00154","title":"Xenbase: Facilitating the Use of Xenopus to Model Human Disease","year":2019,"lang":"en","type":"review","venue":"Frontiers in Physiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":88,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; National Institute of Child Health and Human Development","keywords":"Xenopus; Computational biology; Disease; Biology; Model organism; Bioinformatics; Resource (disambiguation); OMIM : Online Mendelian Inheritance in Man; Computer science; Gene; Genetics; Medicine; Pathology; Phenotype","score_opus":0.10041683964040234,"score_gpt":0.3493090097045996,"score_spread":0.24889217006419728,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2916913049","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015521465,0.019654784,0.03552583,0.0030809506,0.0011121975,0.0013534745,0.84125155,0.05582671,0.026673006],"genre_scores_gemma":[0.030492738,0.020504534,0.16243902,0.0016299967,0.00035425829,0.0021582418,0.7665547,0.010225111,0.0056415023],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980242,0.0005382932,0.0005926717,0.000301644,0.00042242455,0.000120640274],"domain_scores_gemma":[0.99245244,0.004022463,0.001514235,0.0007563162,0.0007554663,0.00049902854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004882611,0.0016452319,0.0014847972,0.013842529,0.0008399372,0.0039314874,0.001999434,0.0011073933,0.029294245],"category_scores_gemma":[0.012240415,0.0007995036,0.002413782,0.008960966,0.00046285806,0.0028722454,0.004133727,0.0011934049,0.017203284],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030868496,0.00027742115,0.014364058,0.058949694,0.0015765394,0.004562616,0.0041029253,0.0045062746,0.045383867,0.025267346,0.682488,0.1554343],"study_design_scores_gemma":[0.00030954974,0.00018402588,0.008949545,0.0059826365,0.0007688431,0.0015851485,0.00043035435,0.002309686,0.0074716904,0.0046221833,0.96723986,0.00014644618],"about_ca_topic_score_codex":0.0036737497,"about_ca_topic_score_gemma":0.0058552404,"teacher_disagreement_score":0.029294245,"about_ca_system_score_codex":0.0011372172,"about_ca_system_score_gemma":0.0033463973,"threshold_uncertainty_score":0.09799904},"labels":[],"label_agreement":null},{"id":"W2917277095","doi":"10.1093/jamia/ocy189","title":"deepBioWSD: effective deep neural word sense disambiguation of biomedical text data","year":2018,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":39,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Artificial intelligence; Natural language processing; Initialization; Pipeline (software); Language model; Machine learning; Vocabulary; Word-sense disambiguation; Deep learning; Word (group theory); Artificial neural network; Lexicon; Unified Medical Language System; Classifier (UML)","score_opus":0.011372290631705834,"score_gpt":0.30435754034277795,"score_spread":0.2929852497110721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917277095","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17357253,0.009199153,0.7341462,0.0045436122,0.0018100425,0.0006663987,0.030711997,0.037683282,0.0076668453],"genre_scores_gemma":[0.47708204,0.0025043755,0.45521098,0.002395211,0.00036429343,0.0006281326,0.050677188,0.0006466548,0.01049122],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993649,0.00010978229,0.00007887582,0.00022499543,0.00016766386,0.00005386183],"domain_scores_gemma":[0.9992324,0.00031153593,0.0001239702,0.00015455244,0.00012482163,0.000052633644],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001191108,0.0014151211,0.0006841019,0.0019500426,0.00060267886,0.0012003033,0.0019674823,0.0011399965,0.0024603023],"category_scores_gemma":[0.003246965,0.0004811062,0.0011447748,0.0015320355,0.0007574344,0.0023176486,0.0024001913,0.0018495473,0.0016898907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00086311897,0.0005066109,0.008128297,0.0013222242,0.00039097303,0.0008402063,0.0005960848,0.11071963,0.0217737,0.009623638,0.0801665,0.7650689],"study_design_scores_gemma":[0.00015504113,0.0001811398,0.0023990374,0.00012558908,0.000085837775,0.00027874316,0.00025737152,0.9171486,0.024059743,0.031738594,0.023501657,0.00006864252],"about_ca_topic_score_codex":0.007474564,"about_ca_topic_score_gemma":0.011462392,"teacher_disagreement_score":0.007474564,"about_ca_system_score_codex":0.001069241,"about_ca_system_score_gemma":0.0018223616,"threshold_uncertainty_score":0.01486212},"labels":[],"label_agreement":null},{"id":"W2917346689","doi":"10.1136/jclinpath-2019-205725","title":"The adaptation of AABACUS for quality improvement in laboratory workflow analysis (\"L-AABACUS\")","year":2019,"lang":"en","type":"article","venue":"Journal of Clinical Pathology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University Health Network","funders":"","keywords":"Workflow; Workload; Computer science; Turnaround time; Staffing; Adaptation (eye); Quality management; Key (lock); Medicine; Operations management; Engineering; Database; Operating system; Management system","score_opus":0.0535430150716385,"score_gpt":0.414244089975711,"score_spread":0.3607010749040725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917346689","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019569527,0.00041708068,0.95794636,0.0010076208,0.0003193102,0.0011494285,0.00031797073,0.015365007,0.0039076502],"genre_scores_gemma":[0.15317363,0.00024924063,0.8427856,0.0006679152,0.000112364905,0.0010248099,0.0005572668,0.0005566637,0.00087249547],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.973196,0.011681453,0.0020583672,0.0026614477,0.0095288055,0.0008739726],"domain_scores_gemma":[0.95958155,0.014530143,0.0046806466,0.007856337,0.011844447,0.00150694],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02133667,0.0014032798,0.000964832,0.0021443367,0.00086892606,0.0038704774,0.002512005,0.0012153103,0.0013457079],"category_scores_gemma":[0.049843304,0.00086431665,0.0015142579,0.0016784574,0.001563161,0.0028378223,0.0045234757,0.0029515238,0.0012578879],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009735945,0.00084878615,0.029168224,0.0014356301,0.00031500778,0.00034495725,0.0017473408,0.070388466,0.04448901,0.055981994,0.017194845,0.77711207],"study_design_scores_gemma":[0.00021249767,0.0018429882,0.016297486,0.0007546245,0.00018165249,0.00073728315,0.00039672502,0.7391426,0.081887744,0.04087135,0.11715216,0.0005228437],"about_ca_topic_score_codex":0.0040093353,"about_ca_topic_score_gemma":0.0025669327,"teacher_disagreement_score":0.02133667,"about_ca_system_score_codex":0.0022844262,"about_ca_system_score_gemma":0.0054005366,"threshold_uncertainty_score":0.11284047},"labels":[],"label_agreement":null},{"id":"W2917538488","doi":"10.7554/elife.43514","title":"Hypothesis, analysis and synthesis, it's all Greek to me","year":2019,"lang":"en","type":"article","venue":"eLife","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Crete; University of Toronto; University of Cyprus","keywords":"Coining (mint); Ancient Greek; Linguistics; Modern Greek; Greek language; Meaning (existential); Computer science; History; Epistemology; Philosophy; Archaeology","score_opus":0.017127236033922758,"score_gpt":0.26416730169163033,"score_spread":0.24704006565770759,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2917538488","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018229919,0.09700926,0.3188556,0.39977154,0.103112526,0.0029921304,0.014990866,0.0022709838,0.042767096],"genre_scores_gemma":[0.30689216,0.049736686,0.45571125,0.108635716,0.029003663,0.01192397,0.0058484375,0.0014328734,0.030815205],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98077846,0.012816781,0.0015302347,0.0023604725,0.0022354044,0.00027874086],"domain_scores_gemma":[0.9174752,0.057411097,0.004702501,0.011372735,0.00806912,0.0009692783],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029357571,0.0012140648,0.0021592325,0.0042978325,0.0014978824,0.0065907924,0.0016022773,0.0022210293,0.02781995],"category_scores_gemma":[0.14034769,0.0005087599,0.0021768494,0.0039621554,0.004978269,0.0069858045,0.0027146665,0.00370714,0.005877794],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011724348,0.00011007079,0.0042765806,0.02400092,0.0024472731,0.0007089653,0.008269131,0.0015987244,0.0019751883,0.3718699,0.2930373,0.2905335],"study_design_scores_gemma":[0.00023619598,0.00023381697,0.0025786536,0.010236007,0.0011275912,0.00035769166,0.0036484916,0.0022263553,0.0012419737,0.55602044,0.42197317,0.00011966115],"about_ca_topic_score_codex":0.0019583064,"about_ca_topic_score_gemma":0.0014985622,"teacher_disagreement_score":0.029357571,"about_ca_system_score_codex":0.0021286712,"about_ca_system_score_gemma":0.0068828748,"threshold_uncertainty_score":0.15525949},"labels":[],"label_agreement":null},{"id":"W2918116898","doi":"10.1186/2041-1480-3-s1-i1","title":"Selected papers from the 14th Annual Bio-Ontologies Special Interest Group Meeting","year":2012,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Biomedicine; Ontology; Computer science; Open Biomedical Ontologies; Presentation (obstetrics); World Wide Web; Data science; Information retrieval; Semantic Web; Upper ontology; Bioinformatics; Ontology alignment; Medicine","score_opus":0.018273993403541747,"score_gpt":0.26382126938263306,"score_spread":0.24554727597909132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2918116898","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003549503,0.079396494,0.009381852,0.10926237,0.65926474,0.0018426033,0.010104574,0.000845752,0.12635209],"genre_scores_gemma":[0.008400456,0.075690664,0.0076810336,0.024399789,0.21480128,0.001330317,0.016736314,0.0010221784,0.649938],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.996999,0.0003665259,0.00025610896,0.00033643836,0.0017429915,0.0002989561],"domain_scores_gemma":[0.9904483,0.0008540595,0.00045638863,0.00023910159,0.005204307,0.0027978227],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0062356275,0.0022477296,0.0020933822,0.0067617977,0.0030060317,0.0076393066,0.002096219,0.0029506679,0.121772],"category_scores_gemma":[0.0076504736,0.0005953546,0.0016765967,0.0058867782,0.0006322598,0.004229358,0.004312261,0.002407355,0.060970087],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007951214,0.000046168432,0.00016965534,0.00027664308,0.00001117549,0.00009511875,0.00004030152,0.000042357406,0.00044732195,0.00030077447,0.9576191,0.040871862],"study_design_scores_gemma":[0.000013878416,0.000025499185,0.0006255585,0.0002054183,0.000013989627,0.000052553314,0.000106305895,0.00007010435,0.0001539466,0.00041268303,0.99830985,0.000010220496],"about_ca_topic_score_codex":0.0031695429,"about_ca_topic_score_gemma":0.012149865,"teacher_disagreement_score":0.121772,"about_ca_system_score_codex":0.003111933,"about_ca_system_score_gemma":0.004302506,"threshold_uncertainty_score":0.4073679},"labels":[],"label_agreement":null},{"id":"W2920153156","doi":"10.1007/978-0-387-39940-9_628","title":"Biological Metadata Management","year":2009,"lang":"en","type":"book-chapter","venue":"Encyclopedia of Database Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Metadata; Computer science; Metadata management; World Wide Web","score_opus":0.026835841295700005,"score_gpt":0.25913468311431426,"score_spread":0.23229884181861427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2920153156","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0044101025,0.010888721,0.40477583,0.0067027346,0.0036292637,0.001152808,0.096570395,0.074684046,0.3971861],"genre_scores_gemma":[0.026551738,0.014894424,0.26295877,0.0051053287,0.0013307864,0.0011150133,0.27456212,0.009804573,0.40367728],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9985428,0.00014325469,0.00022786566,0.00026203255,0.00072249665,0.00010154268],"domain_scores_gemma":[0.9966383,0.0004631585,0.00016770356,0.0016001356,0.00087876205,0.00025194668],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029373905,0.0010328109,0.0009756144,0.007321735,0.0018155972,0.00747344,0.0028352381,0.0011381843,0.07010321],"category_scores_gemma":[0.0044453084,0.0006524635,0.0010017257,0.008034856,0.0007275396,0.0058935327,0.00411747,0.0017633539,0.0908316],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000100858386,0.0001186991,0.0010413727,0.00065578724,0.000041397685,0.0001781079,0.0003046248,0.0005809823,0.013914278,0.054654203,0.3952479,0.53316176],"study_design_scores_gemma":[0.000007746139,0.000009569247,0.00044468147,0.00011698949,0.000018458193,0.0001662924,0.000064199776,0.00068648695,0.006026722,0.014913451,0.9775259,0.000019578987],"about_ca_topic_score_codex":0.003149045,"about_ca_topic_score_gemma":0.0036838988,"teacher_disagreement_score":0.07010321,"about_ca_system_score_codex":0.0015962534,"about_ca_system_score_gemma":0.0031422458,"threshold_uncertainty_score":0.23451859},"labels":[],"label_agreement":null},{"id":"W2920239411","doi":"10.1007/s11036-019-01237-3","title":"Towards Using Scientific Publications to Automatically Extract Information on Rare Diseases","year":2019,"lang":"en","type":"article","venue":"Mobile Networks and Applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"National Institutes of Health","keywords":"Computer science; Statistic; Population; Information extraction; Scientific literature; Rare disease; Data science; Disease; Information retrieval; Medicine; Pathology; Statistics","score_opus":0.010247011819379993,"score_gpt":0.27342695732118455,"score_spread":0.26317994550180457,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2920239411","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11180722,0.016611626,0.7345341,0.009714915,0.0010766848,0.0014248856,0.08857601,0.018142177,0.018112373],"genre_scores_gemma":[0.12656231,0.0062073935,0.78051466,0.00093229616,0.00042309373,0.0004612162,0.080795236,0.00049742544,0.0036063427],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970721,0.000579799,0.00060938246,0.00067611696,0.00091104774,0.00015151336],"domain_scores_gemma":[0.98668706,0.007363678,0.0016358361,0.0012915463,0.002567873,0.00045390666],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0040701493,0.0013039185,0.0012179781,0.027059227,0.001063373,0.0052997074,0.001278663,0.0018636812,0.0020895465],"category_scores_gemma":[0.015867736,0.00050960685,0.0018430898,0.012787743,0.00068507507,0.0047125705,0.0029828176,0.001687651,0.0038268731],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038396515,0.00075223774,0.053219162,0.0053293062,0.00081859616,0.0018738352,0.0020192985,0.005338783,0.051750507,0.023329778,0.046580575,0.8086039],"study_design_scores_gemma":[0.00024753067,0.00043552415,0.0709427,0.0037506777,0.0028794792,0.005689137,0.0043721437,0.16500929,0.09330125,0.14868248,0.50435716,0.00033266368],"about_ca_topic_score_codex":0.0042920276,"about_ca_topic_score_gemma":0.0062675737,"teacher_disagreement_score":0.9729408,"about_ca_system_score_codex":0.0010416519,"about_ca_system_score_gemma":0.0047392473,"threshold_uncertainty_score":0.021525264},"labels":[],"label_agreement":null},{"id":"W2920250724","doi":"10.1007/978-3-030-17083-7_2","title":"Identifying Clinical Terms in Free-Text Notes Using Ontology-Guided Machine Learning","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; Hospital for Sick Children; University of Toronto","funders":"","keywords":"Computer science; Ontology; Artificial intelligence; Natural language processing; Information retrieval","score_opus":0.06266941943699998,"score_gpt":0.3460765542345941,"score_spread":0.2834071347975941,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2920250724","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16463134,0.006140471,0.7477534,0.0032109444,0.00091114454,0.0014628288,0.03928729,0.014602208,0.022000408],"genre_scores_gemma":[0.23077437,0.0019385393,0.7270851,0.0004728404,0.00019379726,0.00034306225,0.031288862,0.00053233217,0.0073710685],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99896944,0.00017312933,0.00017017902,0.00028412158,0.00032692525,0.00007628992],"domain_scores_gemma":[0.99597955,0.0027760167,0.0003416476,0.00025916132,0.00052431517,0.000119208125],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010082844,0.0009380012,0.00054317527,0.004830515,0.00070884207,0.0020158545,0.0011157687,0.0011111574,0.0055272304],"category_scores_gemma":[0.0057832533,0.00029147536,0.0010822512,0.0032429395,0.00043918265,0.0021361103,0.0015959904,0.0012485208,0.0034143622],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043106094,0.0003070496,0.014655109,0.0012589261,0.00012285025,0.0012721563,0.00076312013,0.006036573,0.04177467,0.0064686746,0.034581516,0.8923282],"study_design_scores_gemma":[0.00029058685,0.00052452646,0.051610507,0.0015080819,0.0008094421,0.008488925,0.0037995158,0.53072745,0.12603407,0.078434736,0.19747947,0.00029274783],"about_ca_topic_score_codex":0.0058033955,"about_ca_topic_score_gemma":0.012403979,"teacher_disagreement_score":0.0058033955,"about_ca_system_score_codex":0.00096259214,"about_ca_system_score_gemma":0.002189747,"threshold_uncertainty_score":0.018490434},"labels":[],"label_agreement":null},{"id":"W2921069034","doi":"10.1016/j.prevetmed.2019.03.002","title":"Drivers for the development of an Animal Health Surveillance Ontology (AHSO)","year":2019,"lang":"en","type":"article","venue":"Preventive Veterinary Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Prince Edward Island","funders":"VINNOVA","keywords":"Interoperability; Ontology; Computer science; Animal health; Field (mathematics); Domain (mathematical analysis); Knowledge management; Data science; Component (thermodynamics); Semantic Web; Set (abstract data type); Semantic interoperability; World Wide Web; Medicine","score_opus":0.043119948960159776,"score_gpt":0.3564570117521965,"score_spread":0.31333706279203677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2921069034","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019315831,0.0014981136,0.8687572,0.06774322,0.00094415026,0.0014936428,0.0022862498,0.0022553569,0.03570626],"genre_scores_gemma":[0.06888473,0.0015936235,0.91226345,0.004419068,0.0002072016,0.000725922,0.0058441223,0.00052266655,0.005539231],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9874057,0.0038262324,0.0020204543,0.0014988359,0.004426952,0.00082175067],"domain_scores_gemma":[0.95872533,0.010533675,0.0035917615,0.005373685,0.01800415,0.0037714248],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.044200655,0.0007863284,0.0006269538,0.0047537377,0.0025365008,0.006921964,0.0029300372,0.0027401755,0.0035904895],"category_scores_gemma":[0.03419477,0.00095179846,0.0020167811,0.0041793673,0.0036223927,0.015045553,0.009084634,0.006458531,0.0018726063],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010543267,0.0004089601,0.010996793,0.001356107,0.00013307847,0.0005539148,0.004645849,0.003484512,0.009720841,0.7532952,0.029942991,0.18535632],"study_design_scores_gemma":[0.000038472823,0.00012579374,0.0064099706,0.00187185,0.00015344184,0.0006727891,0.0052154358,0.018113626,0.0060471757,0.306784,0.6544237,0.00014368712],"about_ca_topic_score_codex":0.017841585,"about_ca_topic_score_gemma":0.013446486,"teacher_disagreement_score":0.044200655,"about_ca_system_score_codex":0.005641446,"about_ca_system_score_gemma":0.02007995,"threshold_uncertainty_score":0.23375821},"labels":[],"label_agreement":null},{"id":"W2922023155","doi":"10.1093/jcag/gwz006.208","title":"A209 VALIDATION OF A NATURAL LANGUAGE PROCESSING ALGORITHM TO EXTRACT DATA FOR SYSTEM-LEVEL ADENOMA DETECTION RATE CALCULATION","year":2019,"lang":"en","type":"article","venue":"Journal of the Canadian Association of Gastroenterology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Health Sciences Centre; Sunnybrook Health Science Centre; Cancer Care Ontario; St. Joseph’s Healthcare Hamilton","funders":"","keywords":"Medicine; Algorithm; Test set; Artificial intelligence; Adenoma; Natural language processing; Colorectal cancer; Machine learning; Radiology; Cancer; Internal medicine; Computer science","score_opus":0.014599180450104658,"score_gpt":0.2653532307333676,"score_spread":0.2507540502832629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2922023155","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31319627,0.0003342842,0.6647603,0.00048374577,0.00012181396,0.0035722915,0.007489019,0.007215972,0.002826299],"genre_scores_gemma":[0.30543798,0.00009855928,0.6818336,0.00016915564,0.000033599048,0.0022315062,0.008961075,0.00022263416,0.0010118157],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9861936,0.004868064,0.002497544,0.0030011542,0.0031324653,0.00030719378],"domain_scores_gemma":[0.9421883,0.03659379,0.0024834468,0.0034237492,0.014999299,0.00031132664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019551264,0.0008701054,0.0007324667,0.0027136172,0.00087295583,0.001716488,0.0011023858,0.00089603255,0.0018281589],"category_scores_gemma":[0.069526404,0.00043299727,0.0011209105,0.0016639719,0.0006816802,0.0010321524,0.0012257758,0.00082329,0.001159534],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030182553,0.0020629533,0.15167607,0.0014865061,0.00082553196,0.0011419056,0.002606353,0.07037878,0.088488415,0.003936136,0.010291823,0.66408724],"study_design_scores_gemma":[0.0007930464,0.0018364997,0.10163901,0.00024743614,0.00031087123,0.0012845016,0.0005940501,0.79988754,0.07710593,0.0024089727,0.013713476,0.00017867799],"about_ca_topic_score_codex":0.01685196,"about_ca_topic_score_gemma":0.011920418,"teacher_disagreement_score":0.019551264,"about_ca_system_score_codex":0.0014773858,"about_ca_system_score_gemma":0.003814829,"threshold_uncertainty_score":0.103398204},"labels":[],"label_agreement":null},{"id":"W2928818852","doi":"10.2196/12596","title":"Identifying Clinical Terms in Medical Text Using Ontology-Guided Machine Learning","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":69,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children; University of Toronto","funders":"","keywords":"Computer science; Ontology; Natural language processing; Unified Medical Language System; Artificial intelligence; Information retrieval","score_opus":0.05108806052573915,"score_gpt":0.3965789011479377,"score_spread":0.34549084062219854,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2928818852","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31652975,0.005936069,0.64679605,0.0033818092,0.00048245612,0.0010294301,0.011352987,0.0086575635,0.0058339606],"genre_scores_gemma":[0.53159237,0.0015420711,0.44669828,0.0007388509,0.00024374873,0.00041821387,0.0169132,0.00017496156,0.0016782965],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986206,0.00037344432,0.00023290476,0.00044583617,0.0002528731,0.000074368596],"domain_scores_gemma":[0.99520725,0.0032207153,0.0005771669,0.00030420473,0.0006016864,0.000088949484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013768951,0.0010189211,0.0005097521,0.00549286,0.00048601697,0.0010259618,0.0010333542,0.0009392129,0.0011331857],"category_scores_gemma":[0.0072446098,0.0002391267,0.00097089505,0.0027182086,0.00064859854,0.0022964801,0.0010788635,0.0010800615,0.00078898435],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039875676,0.0005494474,0.027000418,0.001318412,0.00022771051,0.0012270134,0.0009394569,0.06680206,0.022501731,0.0056511555,0.020603709,0.8527801],"study_design_scores_gemma":[0.00008593779,0.00014579542,0.008182958,0.00020197763,0.00015629263,0.00073743856,0.00045178132,0.94388455,0.011757977,0.020030638,0.014308448,0.000056263576],"about_ca_topic_score_codex":0.010824809,"about_ca_topic_score_gemma":0.012920145,"teacher_disagreement_score":0.010824809,"about_ca_system_score_codex":0.0014075361,"about_ca_system_score_gemma":0.002081331,"threshold_uncertainty_score":0.021523595},"labels":[],"label_agreement":null},{"id":"W2929951914","doi":"10.24963/ijcai.2019/845","title":"EL Embeddings: Geometric Construction of Models for the Description Logic EL++","year":2019,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Embedding; Vector space; Interpretation (philosophy); Similarity (geometry); Theoretical computer science; Statistical relational learning; Semantics (computer science); Computer science; Space (punctuation); Algebraic number; Mathematics; Discrete mathematics; Artificial intelligence; Relational database; Pure mathematics; Data mining; Image (mathematics)","score_opus":0.050129356721033226,"score_gpt":0.3085492676816046,"score_spread":0.2584199109605714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2929951914","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049966946,0.00017312041,0.99062365,0.00036915857,0.000045502384,0.000049010338,0.00030651185,0.00047696393,0.0029593161],"genre_scores_gemma":[0.21358335,0.0006097742,0.77401847,0.00046828113,0.00012770011,0.0003498074,0.0024598802,0.0007122371,0.007670437],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9981951,0.0006947075,0.00017071606,0.00035321648,0.00045615327,0.000130105],"domain_scores_gemma":[0.99810934,0.0008245273,0.00022156778,0.00048148792,0.00024239057,0.000120690616],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019779124,0.0011546303,0.0006448163,0.0023187955,0.00080601487,0.0037802786,0.0018275793,0.0014301152,0.005216601],"category_scores_gemma":[0.006379773,0.0009307128,0.0022515678,0.0021603145,0.0025219542,0.0077904826,0.0065157237,0.0034188644,0.0017553889],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000050324805,0.0000336135,0.00032283075,0.00014619724,0.000026379314,0.00008749974,0.00032340098,0.01811766,0.0010231386,0.9142957,0.0038182999,0.061755],"study_design_scores_gemma":[0.000017526514,0.000030815216,0.00008493575,0.000047290687,0.000014523879,0.000100635414,0.00012392316,0.09961978,0.0010923464,0.88270473,0.016143419,0.000020031737],"about_ca_topic_score_codex":0.0015077076,"about_ca_topic_score_gemma":0.0017749438,"teacher_disagreement_score":0.005216601,"about_ca_system_score_codex":0.0017374934,"about_ca_system_score_gemma":0.00086518674,"threshold_uncertainty_score":0.017451286},"labels":[],"label_agreement":null},{"id":"W2936715908","doi":"10.3389/frma.2019.00002","title":"Editorial: Mining Scientific Papers: NLP-enhanced Bibliometrics","year":2019,"lang":"en","type":"editorial","venue":"Frontiers in Research Metrics and Analytics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Deutscher Akademischer Austauschdienst; Department of Science and Technology, Ministry of Science and Technology, India","keywords":"Bibliometrics; Library science; Front (military); Geography; Data science; History; Computer science","score_opus":0.07409031425835062,"score_gpt":0.4044010460789305,"score_spread":0.3303107318205799,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2936715908","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000044870772,0.0028924856,0.00023668128,0.03484983,0.96052724,0.000040671664,0.00017904358,0.00008940503,0.0011397353],"genre_scores_gemma":[0.00042127367,0.0019858724,0.00021317495,0.010995561,0.9795212,0.000046039444,0.00009737093,0.000073582334,0.0066458825],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9887532,0.0020642385,0.0015871624,0.0013316887,0.00576143,0.0005022509],"domain_scores_gemma":[0.9367782,0.028008513,0.0039014386,0.00115088,0.023118,0.0070430124],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.015341084,0.0042908047,0.005081235,0.011947558,0.0042492934,0.015058923,0.0039069196,0.01527939,0.021441394],"category_scores_gemma":[0.060876004,0.0015145954,0.003928027,0.0039333478,0.0030916638,0.0056962925,0.0028230061,0.015997846,0.015341833],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047610538,0.000013298333,0.00002687616,0.00021042123,0.000034467932,0.000083174134,0.000009084187,0.000025282197,0.0000611784,0.0001583666,0.99631715,0.0030132039],"study_design_scores_gemma":[0.0002502267,0.000049445032,0.00064341695,0.001049377,0.00019415014,0.00042544605,0.00006432467,0.00059487787,0.0003028858,0.0032564597,0.99311745,0.000051926636],"about_ca_topic_score_codex":0.0015061679,"about_ca_topic_score_gemma":0.004638265,"teacher_disagreement_score":0.9880524,"about_ca_system_score_codex":0.0033872263,"about_ca_system_score_gemma":0.0042616734,"threshold_uncertainty_score":0.08113241},"labels":[],"label_agreement":null},{"id":"W2941911965","doi":"10.1007/978-3-030-19432-1_8","title":"A Logical Framework for Modelling Breast Cancer Progression","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science; Mathematical proof; Breast cancer; Cancer; Automated theorem proving; Machine learning; Theoretical computer science; Artificial intelligence; Medicine; Internal medicine","score_opus":0.029641480201894825,"score_gpt":0.3124779018860169,"score_spread":0.28283642168412204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2941911965","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0040432783,0.00056691346,0.98673147,0.00086088735,0.00009181275,0.00012053949,0.0019694187,0.0014906205,0.00412506],"genre_scores_gemma":[0.105850466,0.0010557756,0.8839016,0.0002870545,0.000093505616,0.00038711895,0.003124671,0.00029966928,0.0050002337],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999156,0.00024872413,0.00012108409,0.00018653733,0.00021101556,0.000076502285],"domain_scores_gemma":[0.9986468,0.00082192034,0.000109319175,0.00014135476,0.00021374528,0.00006691373],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018435542,0.0008589019,0.0006253282,0.0013756688,0.00084016775,0.0035799607,0.0024599254,0.0012062386,0.0060934513],"category_scores_gemma":[0.0038584743,0.0006827807,0.0025416997,0.0013982967,0.0013894672,0.0036097325,0.0020119643,0.0014865941,0.0012369179],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001568261,0.00009092045,0.0018243209,0.0004714512,0.00014757847,0.00068076735,0.0006202162,0.3207107,0.0022744597,0.5773109,0.008935398,0.08677641],"study_design_scores_gemma":[0.000041761356,0.00003929658,0.00023030999,0.000092142756,0.00016147146,0.00017620511,0.00012670053,0.61306196,0.0017900085,0.3358053,0.048441883,0.000032906155],"about_ca_topic_score_codex":0.019562628,"about_ca_topic_score_gemma":0.023463696,"teacher_disagreement_score":0.019562628,"about_ca_system_score_codex":0.0017517237,"about_ca_system_score_gemma":0.0023678513,"threshold_uncertainty_score":0.038897514},"labels":[],"label_agreement":null},{"id":"W2945539778","doi":"10.2196/13590","title":"Fast Healthcare Interoperability Resources, Clinical Quality Language, and Systematized Nomenclature of Medicine—Clinical Terms in Representing Clinical Evidence Logic Statements for the Use of Imaging Procedures: Descriptive Study","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Interoperability; Evidence-based medicine; Computer science; Quality (philosophy); Health care; Medicine; Medical physics; Medical education; Alternative medicine; Pathology; World Wide Web","score_opus":0.216553290292055,"score_gpt":0.5207034948081783,"score_spread":0.30415020451612335,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945539778","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8242342,0.050338287,0.03968374,0.0066432124,0.00029310136,0.009250204,0.023856072,0.00014166671,0.04555954],"genre_scores_gemma":[0.91393477,0.016574869,0.04564104,0.0021602742,0.00010836373,0.0070445207,0.012489052,0.00009537255,0.0019516497],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9605634,0.018367294,0.010022337,0.001999511,0.007882262,0.0011651432],"domain_scores_gemma":[0.8603443,0.09096405,0.02405865,0.005517062,0.018055117,0.0010608473],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.03450854,0.00033265405,0.0006158231,0.013618286,0.0011919758,0.0047568586,0.0010646589,0.0008285689,0.0037960855],"category_scores_gemma":[0.12533908,0.00041313568,0.0012993935,0.023717456,0.002449122,0.005995625,0.0033812227,0.0013920319,0.00072819926],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010000374,0.0007728567,0.5926459,0.012670764,0.00043807155,0.0012726295,0.05477872,0.0007371805,0.00088199135,0.035186045,0.013863235,0.28575256],"study_design_scores_gemma":[0.00021090852,0.001201123,0.5562424,0.03406455,0.001932826,0.009522139,0.13999917,0.004519993,0.0028702726,0.016706076,0.23242605,0.000304516],"about_ca_topic_score_codex":0.016384188,"about_ca_topic_score_gemma":0.018786334,"teacher_disagreement_score":0.99524313,"about_ca_system_score_codex":0.0067998823,"about_ca_system_score_gemma":0.014411105,"threshold_uncertainty_score":0.18250078},"labels":[],"label_agreement":null},{"id":"W2946021593","doi":"10.1177/0840470419845384","title":"Organizational implications of implementing a new adverse drug event reporting system for care providers and integrating it with provincial health information systems","year":2019,"lang":"en","type":"article","venue":"Healthcare Management Forum","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vancouver General Hospital; University of British Columbia; Vancouver Coastal Health Research Institute; Simon Fraser University; Vancouver Coastal Health","funders":"Michael Smith Health Research BC","keywords":"Interoperability; Business; Negotiation; Government (linguistics); Health care; Knowledge management; Stakeholder; Event (particle physics); Process management; Information system; Public relations; Computer science; Engineering; Political science","score_opus":0.010763948081156273,"score_gpt":0.2860297811891451,"score_spread":0.27526583310798886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946021593","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5845377,0.0006113274,0.052269995,0.29859397,0.00044963113,0.0023774758,0.0002921997,0.00061850535,0.060249165],"genre_scores_gemma":[0.9541235,0.00026376214,0.03792874,0.0048734364,0.00005678367,0.00032186034,0.00015358781,0.00005570736,0.002222627],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.8853799,0.07942778,0.006211521,0.0031748682,0.013088316,0.01271764],"domain_scores_gemma":[0.8659148,0.057647187,0.01197613,0.014206714,0.030013153,0.020241972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08062952,0.00038430287,0.00030670824,0.0015378497,0.014866708,0.01752001,0.0034440148,0.0041837436,0.002728843],"category_scores_gemma":[0.12696353,0.0006778971,0.00088289066,0.003278715,0.007896395,0.00788545,0.01111698,0.004597528,0.000470474],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00079278223,0.0037003467,0.26212403,0.0015217834,0.00026456246,0.006296524,0.1444045,0.028174859,0.017051958,0.21615142,0.026237937,0.2932793],"study_design_scores_gemma":[0.0005121287,0.0024641948,0.14147177,0.0017327857,0.00037789793,0.002391383,0.4967032,0.044096265,0.012223081,0.08149076,0.2159977,0.0005388672],"about_ca_topic_score_codex":0.1631458,"about_ca_topic_score_gemma":0.15276943,"teacher_disagreement_score":0.1631458,"about_ca_system_score_codex":0.041893493,"about_ca_system_score_gemma":0.13586153,"threshold_uncertainty_score":0.42641473},"labels":[],"label_agreement":null},{"id":"W2946325139","doi":"10.1007/978-3-030-18305-9_54","title":"Semantic Roles: Towards Rhetorical Moves in Writing About Experimental Procedures","year":2019,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; University of Waterloo","funders":"","keywords":"Rhetorical question; Computer science; FrameNet; Natural language processing; Linguistics; Artificial intelligence; Focus (optics); Task (project management); Domain (mathematical analysis); Parsing","score_opus":0.016847739278021037,"score_gpt":0.28526324052888435,"score_spread":0.26841550125086333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2946325139","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0031179294,0.0010127747,0.9136295,0.014161463,0.0010317864,0.00014668278,0.00033697105,0.0014351306,0.06512769],"genre_scores_gemma":[0.15079676,0.0013183258,0.8111323,0.0024791982,0.0010822457,0.0005033908,0.0011534286,0.0018879587,0.029646358],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98697144,0.009306473,0.00070198573,0.0010965593,0.0016191344,0.00030441443],"domain_scores_gemma":[0.9746513,0.02063026,0.000619947,0.0018187923,0.0017697554,0.00050992484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0113159465,0.0015769986,0.0008640494,0.0034045032,0.004133401,0.010313619,0.0039536715,0.0044986256,0.015657915],"category_scores_gemma":[0.026562132,0.0019031133,0.0014887538,0.0023084537,0.012475398,0.031596247,0.0074277185,0.0083764605,0.006685422],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002308759,0.000013371302,0.000035726545,0.00015540097,0.0000052541745,0.000046360085,0.0039297407,0.0002292909,0.00062327617,0.9674444,0.007938915,0.01955509],"study_design_scores_gemma":[0.000023897313,0.000014767411,0.000040515588,0.00022523703,0.000024340356,0.00009622098,0.0016451236,0.005339641,0.0023342466,0.8410244,0.1492091,0.00002250748],"about_ca_topic_score_codex":0.0015040297,"about_ca_topic_score_gemma":0.001736712,"teacher_disagreement_score":0.015657915,"about_ca_system_score_codex":0.0029876528,"about_ca_system_score_gemma":0.0029600088,"threshold_uncertainty_score":0.05984515},"labels":[],"label_agreement":null},{"id":"W2947100013","doi":"10.1186/s12859-019-2801-x","title":"Automated assessment of biological database assertions using the scientific literature","year":2019,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Bhabha Atomic Research Centre; Australian Research Council","keywords":"Computer science; Consistency (knowledge bases); Information retrieval; Assertion; Relevance (law); Relation (database); Biological database; Classifier (UML); Set (abstract data type); Data mining; Data science; Database; Bioinformatics; Artificial intelligence; Biology; Programming language","score_opus":0.045762540862951104,"score_gpt":0.34215141589784726,"score_spread":0.29638887503489614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2947100013","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.50500166,0.011225694,0.40234768,0.0027922126,0.00035908344,0.0018296165,0.029464314,0.038053308,0.00892653],"genre_scores_gemma":[0.5306989,0.001196376,0.43341228,0.00030196737,0.00023617731,0.00033015865,0.0323475,0.00040960815,0.0010670853],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9840577,0.0036583606,0.0026407943,0.0023896564,0.0069019473,0.00035151135],"domain_scores_gemma":[0.86854345,0.07798688,0.01400302,0.008757228,0.028889097,0.0018202905],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01567042,0.0012368548,0.0011707848,0.042697478,0.0014147358,0.0045163115,0.0023497853,0.0016203546,0.0019042362],"category_scores_gemma":[0.08046678,0.00038455613,0.0011457053,0.01148153,0.000707281,0.0034182211,0.003459313,0.0008554673,0.001377871],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013086902,0.00055315305,0.16566287,0.0061188675,0.0008405885,0.0020614236,0.0034323542,0.009749067,0.045019522,0.0057480126,0.029455055,0.7300504],"study_design_scores_gemma":[0.00029012802,0.0009260915,0.19031313,0.0019086681,0.0016169924,0.0061134077,0.0047818352,0.56943035,0.123670794,0.019129219,0.08140146,0.0004178644],"about_ca_topic_score_codex":0.006640622,"about_ca_topic_score_gemma":0.008072565,"teacher_disagreement_score":0.9843296,"about_ca_system_score_codex":0.0013328494,"about_ca_system_score_gemma":0.0041124127,"threshold_uncertainty_score":0.08287406},"labels":[],"label_agreement":null},{"id":"W2947689917","doi":"","title":"KlickLabs at TREC 2018 Precision Medicine track.","year":2018,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Track (disk drive); Precision medicine; Information retrieval; Artificial intelligence; Medicine","score_opus":0.04302524158150629,"score_gpt":0.3146855815370476,"score_spread":0.27166033995554134,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2947689917","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042078914,0.010719389,0.013708896,0.024400113,0.0117534315,0.0010478322,0.8263688,0.04214679,0.06564679],"genre_scores_gemma":[0.005906092,0.0021198003,0.014939035,0.0021801032,0.0012106627,0.0005395524,0.90350807,0.0022176176,0.06737916],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941239,0.0012004739,0.00045471673,0.001026046,0.0024815663,0.00071326894],"domain_scores_gemma":[0.98046976,0.0035418293,0.0007079813,0.0024602234,0.010314412,0.0025058095],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009426671,0.004717301,0.0037186334,0.011090413,0.0034558927,0.0088870125,0.0053260303,0.0034761967,0.18356618],"category_scores_gemma":[0.022717811,0.0010756422,0.001986421,0.008555487,0.0010773224,0.0111951735,0.0048800493,0.003768427,0.18293713],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000056383833,0.00004581788,0.00008601863,0.00017431875,0.000016081056,0.000018451552,0.00001116886,0.00011755513,0.00029780914,0.00029779464,0.9896768,0.009201858],"study_design_scores_gemma":[0.00041397213,0.00015400692,0.0030309055,0.0003752885,0.0001276161,0.00017025162,0.0002420317,0.005810654,0.0044294796,0.006621255,0.97850966,0.00011494044],"about_ca_topic_score_codex":0.07185958,"about_ca_topic_score_gemma":0.112798765,"teacher_disagreement_score":0.18356618,"about_ca_system_score_codex":0.006463628,"about_ca_system_score_gemma":0.008451271,"threshold_uncertainty_score":0.6140901},"labels":[],"label_agreement":null},{"id":"W2949176808","doi":"10.1093/bioinformatics/bty449","title":"Transfer learning for biomedical named entity recognition with neural networks","year":2018,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":195,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Human Genome Research Institute; National Institutes of Health; Nvidia","keywords":"Computer science; Named-entity recognition; Transfer of learning; Artificial intelligence; Annotation; Deep learning; Conditional random field; Task (project management); Artificial neural network; Field (mathematics); Machine learning; Natural language processing; Word (group theory); Noise (video); State (computer science); Domain (mathematical analysis); Mathematics","score_opus":0.019899863939700585,"score_gpt":0.25561859100615614,"score_spread":0.23571872706645555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949176808","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09474929,0.006156334,0.86162174,0.0022077241,0.0006036463,0.00044251143,0.0026584584,0.024155933,0.007404307],"genre_scores_gemma":[0.7199887,0.0018604076,0.25570902,0.00088986964,0.00042583913,0.0006811883,0.011678108,0.00039530132,0.008371527],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983621,0.000484906,0.00012121517,0.00051143556,0.00038368715,0.00013671443],"domain_scores_gemma":[0.9957212,0.0023895197,0.00035749053,0.0006655814,0.0007750975,0.00009103172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003987598,0.0017970223,0.0010270448,0.0024225484,0.0006060783,0.0011241961,0.0026852577,0.0020536869,0.0043917275],"category_scores_gemma":[0.010426974,0.00042943744,0.0009836855,0.002679311,0.001009536,0.0037517084,0.002332265,0.002380607,0.0029152967],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034129867,0.00030606147,0.0021953485,0.00036996728,0.0001844232,0.00028647124,0.00011725605,0.32673857,0.0055353777,0.004560261,0.013359896,0.64600503],"study_design_scores_gemma":[0.000014617402,0.00005804282,0.00060534687,0.00003047781,0.000024617342,0.000053187454,0.000040764957,0.97544724,0.007507289,0.013257396,0.002941881,0.000019125528],"about_ca_topic_score_codex":0.0057705673,"about_ca_topic_score_gemma":0.0043840325,"teacher_disagreement_score":0.0057705673,"about_ca_system_score_codex":0.0017958849,"about_ca_system_score_gemma":0.0011917595,"threshold_uncertainty_score":0.02108866},"labels":[],"label_agreement":null},{"id":"W2949764148","doi":"10.5539/ijsp.v7n4p11","title":"The Application of Text Mining Algorithms In Summarizing Trends in Anti-Epileptic Drug Research","year":2018,"lang":"en","type":"article","venue":"International Journal of Statistics and Probability","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Sentiment analysis; Computer science; Lamotrigine; Topiramate; Latent Dirichlet allocation; Topic model; Information retrieval; Data science; Text mining; tf–idf; Key (lock); Data mining; Machine learning; Epilepsy; Medicine; Psychiatry","score_opus":0.03525369295334435,"score_gpt":0.37503549586323054,"score_spread":0.3397818029098862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949764148","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14920235,0.010458846,0.80385786,0.0030804747,0.00054403953,0.0014545929,0.020087365,0.0069191284,0.004395369],"genre_scores_gemma":[0.29146987,0.004889618,0.6795935,0.00030229523,0.0006937434,0.00083706575,0.02069525,0.00023249951,0.0012861749],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972856,0.00076675235,0.0006792842,0.000567218,0.0006038171,0.00009732714],"domain_scores_gemma":[0.98613495,0.009168845,0.0020599803,0.0006465248,0.0018027521,0.00018708303],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.004871214,0.0014603719,0.0013097036,0.021494774,0.00075155625,0.00306958,0.0008972812,0.0008443949,0.0010104832],"category_scores_gemma":[0.017340634,0.00037644332,0.0017030038,0.014766097,0.00035673912,0.0035922353,0.0008507736,0.000946948,0.0010897835],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035044327,0.00039737247,0.036089025,0.0024240154,0.00081596436,0.00048417193,0.0011725127,0.02546145,0.016844643,0.0061939885,0.01270562,0.8970609],"study_design_scores_gemma":[0.00017830252,0.00085501914,0.06037237,0.00085556076,0.0013907412,0.0011344227,0.0029474862,0.7597309,0.031451914,0.083757795,0.057049062,0.0002764531],"about_ca_topic_score_codex":0.0019038962,"about_ca_topic_score_gemma":0.0022163705,"teacher_disagreement_score":0.97850525,"about_ca_system_score_codex":0.0008261915,"about_ca_system_score_gemma":0.0013765122,"threshold_uncertainty_score":0.025761783},"labels":[],"label_agreement":null},{"id":"W2949926253","doi":"10.1093/bioinformatics/bty722","title":"Adjutant: an R-based tool to support topic discovery for systematic and literature reviews","year":2018,"lang":"en","type":"review","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Centre for Disease Control; University of British Columbia","funders":"Canadian Institutes of Health Research; Canada Research Chairs; Michael Smith Health Research BC","keywords":"Systematic review; Computer science; Data science; Information retrieval; MEDLINE; Biology","score_opus":0.04961544002561298,"score_gpt":0.3508484078449882,"score_spread":0.30123296781937525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2949926253","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043477537,0.013094279,0.20039602,0.0037939865,0.001855098,0.0069285682,0.4873695,0.26679373,0.015421022],"genre_scores_gemma":[0.025056327,0.0076540816,0.7150134,0.0017840421,0.00060020504,0.031068781,0.16325451,0.04496195,0.010606734],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98779374,0.0051862146,0.0027100309,0.0022176637,0.0017004691,0.0003919336],"domain_scores_gemma":[0.8898902,0.08118626,0.011132302,0.0065905494,0.009166508,0.0020341289],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01793446,0.0036433118,0.0052747773,0.024029307,0.0016716914,0.0065428624,0.0042923386,0.0016030448,0.15821916],"category_scores_gemma":[0.11127992,0.0026981374,0.0077365413,0.021539148,0.0013111158,0.0041082627,0.007095552,0.0027058395,0.06761907],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017607469,0.00008256655,0.004868162,0.15736997,0.0053734817,0.0013325934,0.0018905612,0.0030273912,0.008014574,0.015229859,0.65602815,0.14502195],"study_design_scores_gemma":[0.0019881204,0.00030076675,0.00749222,0.014674645,0.005549242,0.0019565546,0.00046376613,0.011742223,0.00853753,0.04060745,0.9060903,0.00059707894],"about_ca_topic_score_codex":0.0021377981,"about_ca_topic_score_gemma":0.006434379,"teacher_disagreement_score":0.98206556,"about_ca_system_score_codex":0.0022830965,"about_ca_system_score_gemma":0.015275668,"threshold_uncertainty_score":0.5292958},"labels":[],"label_agreement":null},{"id":"W2950150962","doi":"10.1038/npre.2010.5382","title":"SPARQL Assist Language Neutral Query Composer","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Paul's Hospital; University of British Columbia","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Technische Universität Dortmund; Canarie; Microsoft Research","keywords":"SPARQL; Computer science; Named graph; XML; Identifier; Task (project management); World Wide Web; Context (archaeology); Semantic Web; Information retrieval; Linked data; RDF; Programming language; Engineering","score_opus":0.007474534731342467,"score_gpt":0.2968057852026288,"score_spread":0.2893312504712863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2950150962","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013378173,0.00018706312,0.70468193,0.0007952265,0.00023565971,0.0006246872,0.0036427605,0.26052925,0.015925176],"genre_scores_gemma":[0.31863904,0.00052036415,0.5586678,0.0028430438,0.0004518845,0.0009768454,0.017665563,0.056132436,0.044103153],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9936625,0.0015226379,0.0007402607,0.0011402991,0.0025460823,0.00038829003],"domain_scores_gemma":[0.9919932,0.0028798208,0.00030850386,0.0030136488,0.0015489466,0.0002559039],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075420635,0.0014874062,0.0011519048,0.0014123665,0.0007724682,0.0031259526,0.0021175214,0.0012724554,0.027594097],"category_scores_gemma":[0.011686636,0.0008263888,0.0010709547,0.0012599665,0.001057214,0.00407595,0.0045835553,0.0017374172,0.012645156],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004075558,0.00056832854,0.007359402,0.0014626583,0.00029508353,0.0031477003,0.0027214836,0.011849688,0.08616423,0.11051959,0.31917754,0.4526587],"study_design_scores_gemma":[0.00039125772,0.00019725911,0.0016216432,0.0001967631,0.00013470196,0.0020541057,0.0006729798,0.19406219,0.1797463,0.06990321,0.5508119,0.00020776945],"about_ca_topic_score_codex":0.0013934857,"about_ca_topic_score_gemma":0.0011838552,"teacher_disagreement_score":0.027594097,"about_ca_system_score_codex":0.0008131043,"about_ca_system_score_gemma":0.0010520547,"threshold_uncertainty_score":0.09231144},"labels":[],"label_agreement":null},{"id":"W2951164180","doi":"10.1101/199570","title":"Evidence-Based Design and Evaluation of a Whole Genome Sequencing Clinical Report for the Reference Microbiology Laboratory","year":2017,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"BC Centre for Disease Control; University of British Columbia","funders":"University of British Columbia","keywords":"Workflow; Computer science; Data science; Evidence-based design; Test (biology); Domain (mathematical analysis); Visualization; Genomics; Benchmarking; Medicine; Data mining; Genome; Pathology","score_opus":0.16647093345321426,"score_gpt":0.35029011889670036,"score_spread":0.1838191854434861,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951164180","genre_codex":"protocol","genre_gemma":"methods","domain_codex":"methods","domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31103957,0.010386567,0.23188555,0.018120395,0.005139375,0.4044968,0.0022502139,0.0019985912,0.014682917],"genre_scores_gemma":[0.32090652,0.0024316127,0.5606496,0.002009392,0.00034239164,0.112253666,0.00063446804,0.00015571056,0.00061672705],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.32388926,0.5891013,0.047555055,0.006592514,0.028793378,0.0040685134],"domain_scores_gemma":[0.23003395,0.5433487,0.05403689,0.043761473,0.11722315,0.01159579],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.5544063,0.0022913935,0.002273769,0.005986453,0.003190928,0.011351178,0.009729284,0.006528905,0.005463079],"category_scores_gemma":[0.69414836,0.0019942315,0.005964056,0.0033584693,0.0035800969,0.006096675,0.008106799,0.0046028993,0.001351673],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.028563099,0.019916296,0.026246263,0.057083573,0.0043645543,0.0015474389,0.022780798,0.02010116,0.002252074,0.01106256,0.013829353,0.7922529],"study_design_scores_gemma":[0.09780621,0.3762605,0.041890256,0.15222043,0.015012902,0.001347585,0.03858251,0.061154768,0.023328723,0.021693476,0.16859934,0.0021033867],"about_ca_topic_score_codex":0.0031825516,"about_ca_topic_score_gemma":0.0044055097,"teacher_disagreement_score":0.5544063,"about_ca_system_score_codex":0.019169932,"about_ca_system_score_gemma":0.049112566,"threshold_uncertainty_score":0.5494964},"labels":[],"label_agreement":null},{"id":"W2951333474","doi":"10.1186/s13326-016-0067-z","title":"FALDO: a semantic standard for describing the location of nucleotide and protein feature annotation","year":2016,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hydro One (Canada)","funders":"National Bioscience Database Center; Basic Energy Sciences; Japan Science and Technology Agency; National Institutes of Health; Research Organization of Information and Systems; Swiss Institute of Bioinformatics; Staatssekretariat für Bildung, Forschung und Innovation; U.S. Department of Energy; Australian Government; Office of Science; Scottish Government","keywords":"SPARQL; UniProt; Computer science; Annotation; Ontology; Information retrieval; Feature (linguistics); Semantic Web; Computational biology; RDF; Biology; Artificial intelligence; Genetics; Gene","score_opus":0.02384984259902257,"score_gpt":0.26560733194212777,"score_spread":0.2417574893431052,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951333474","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008411813,0.00065961434,0.8400033,0.0024847456,0.000538921,0.0010620615,0.07972488,0.04425231,0.022862427],"genre_scores_gemma":[0.07292979,0.0017219491,0.65094364,0.0038995815,0.0002698446,0.0025300698,0.24884482,0.0074359053,0.011424327],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9943084,0.00095003954,0.0020329775,0.0006889983,0.0015855615,0.00043401564],"domain_scores_gemma":[0.98722756,0.0030753934,0.0012161236,0.0044816216,0.0033926493,0.00060657697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008933141,0.0012706865,0.00094824453,0.0075330357,0.0019916792,0.0055610472,0.003270347,0.003156892,0.009122744],"category_scores_gemma":[0.015083088,0.0010394126,0.0016260876,0.0053517246,0.0023793583,0.010266426,0.0054281144,0.0034743126,0.008077729],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009967574,0.0003381772,0.0075674094,0.0022468264,0.00015913691,0.0008645942,0.002634179,0.006854875,0.029615,0.53091526,0.23006664,0.18774112],"study_design_scores_gemma":[0.00012406103,0.0001127135,0.0025490446,0.0008398869,0.00007170877,0.00084378145,0.0007979987,0.0150072845,0.013653711,0.10145208,0.8643513,0.00019637142],"about_ca_topic_score_codex":0.01674253,"about_ca_topic_score_gemma":0.015239703,"teacher_disagreement_score":0.01674253,"about_ca_system_score_codex":0.0035149823,"about_ca_system_score_gemma":0.0057766247,"threshold_uncertainty_score":0.047243536},"labels":[],"label_agreement":null},{"id":"W2951560185","doi":"10.1038/s41592-019-0422-y","title":"CancerMine: a literature-mined resource for drivers, oncogenes and tumor suppressors in cancer","year":2019,"lang":"en","type":"article","venue":"Nature Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":223,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"","keywords":"Suppressor; License; Resource (disambiguation); Cancer; Disease; Database; Computer science; Bioinformatics; Biology; Medicine; Genetics; Internal medicine","score_opus":0.011490399213881225,"score_gpt":0.3900107605720499,"score_spread":0.3785203613581687,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2951560185","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006418672,0.018906461,0.019208403,0.0009551019,0.00022283763,0.0004505468,0.93391234,0.011102789,0.008822817],"genre_scores_gemma":[0.030034687,0.018802674,0.07910238,0.001000054,0.00017887086,0.0013744149,0.86457145,0.0019553965,0.0029801722],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99882406,0.00020436275,0.00032716518,0.00031406962,0.0002515753,0.00007875858],"domain_scores_gemma":[0.99381363,0.00410161,0.0006342993,0.00049871934,0.00059770804,0.0003540823],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019329984,0.0020881959,0.0021640512,0.024437748,0.0014694318,0.0025766029,0.002482893,0.002079279,0.025456581],"category_scores_gemma":[0.012000711,0.00088341406,0.0021388943,0.015619604,0.0005594576,0.0020473367,0.0030649172,0.0015131494,0.012013545],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016244658,0.00032916627,0.016405018,0.09851734,0.0023865046,0.0052758856,0.001955593,0.005766763,0.03609233,0.026872685,0.55782634,0.24694796],"study_design_scores_gemma":[0.00028091975,0.00013545838,0.009514043,0.005710546,0.0024795039,0.0030164612,0.0003987083,0.0041995565,0.009046044,0.011389778,0.95369816,0.00013086775],"about_ca_topic_score_codex":0.0052943183,"about_ca_topic_score_gemma":0.010668438,"teacher_disagreement_score":0.025456581,"about_ca_system_score_codex":0.0011655121,"about_ca_system_score_gemma":0.0054590525,"threshold_uncertainty_score":0.08516073},"labels":[],"label_agreement":null},{"id":"W2953749868","doi":"10.2196/15063","title":"Authorship Correction: A Clinical Decision Support Engine Based on a National Medication Repository for the Detection of Potential Duplicate Medications: Design and Evaluation","year":2019,"lang":"en","type":"erratum","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Clinical decision support system; Computer science; Decision support system; Medicine; Medical emergency; Medical physics; Data mining","score_opus":0.04033442004054544,"score_gpt":0.37123861174861394,"score_spread":0.3309041917080685,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953749868","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022871364,0.0011479964,0.014130939,0.20821507,0.7467991,0.00026103383,0.011915192,0.0041710404,0.011072413],"genre_scores_gemma":[0.10524044,0.0062303892,0.08849047,0.20551684,0.12696826,0.0013195315,0.027089298,0.010152668,0.42899203],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9884532,0.0023407082,0.0021795237,0.0011102941,0.0053369594,0.0005792405],"domain_scores_gemma":[0.8746923,0.04926151,0.003921966,0.006887724,0.062304676,0.0029318433],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009555612,0.0011521818,0.0011040437,0.0022115915,0.0028275857,0.0035733534,0.0027653305,0.0049821255,0.08843932],"category_scores_gemma":[0.16272062,0.0007192101,0.0012398019,0.0020682397,0.0020919074,0.0020117748,0.002692431,0.0049073817,0.028068805],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009290565,0.00001766627,0.0002974115,0.00023523059,0.000022083763,0.0005061135,0.000086225926,0.0001080773,0.00016683814,0.0013287563,0.97787523,0.019263387],"study_design_scores_gemma":[0.00016722093,0.00008593746,0.0011235027,0.00067525177,0.00010950303,0.0016268205,0.00022098146,0.0014264233,0.0017114306,0.0030959535,0.98967874,0.000078186305],"about_ca_topic_score_codex":0.010703019,"about_ca_topic_score_gemma":0.018062932,"teacher_disagreement_score":0.08843932,"about_ca_system_score_codex":0.0035634434,"about_ca_system_score_gemma":0.008935071,"threshold_uncertainty_score":0.29585898},"labels":[],"label_agreement":null},{"id":"W2953936633","doi":"10.7152/acro.v29i1.15463","title":"Examining Communities in the Transdisciplinary Area of Cognitive Science: Automatic Classification for Examining Communities in the Web of Science Using Unsupervised Clustering Methods","year":2019,"lang":"en","type":"article","venue":"Advances in Classification Research Online","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Scopus; Subject (documents); Cluster analysis; Scope (computer science); Domain (mathematical analysis); Data science; Web of science; Information retrieval; Cognition; World Wide Web; Artificial intelligence; MEDLINE; Psychology","score_opus":0.37145671480616316,"score_gpt":0.522537251833423,"score_spread":0.15108053702725988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2953936633","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.118541375,0.0006396733,0.87320834,0.00060423475,0.00005766899,0.0007410889,0.0011727691,0.001457069,0.0035777418],"genre_scores_gemma":[0.2963301,0.00016656844,0.6999575,0.0000774876,0.00004739171,0.0005825187,0.0013954907,0.00014653927,0.0012965066],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949502,0.0013689575,0.00046447795,0.0012890827,0.001574473,0.00035274998],"domain_scores_gemma":[0.9851259,0.0060621058,0.0025408738,0.0016979194,0.0038271728,0.0007461262],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0047077043,0.0006935027,0.00073101907,0.019171178,0.002316154,0.0038022378,0.0014875117,0.0015156463,0.0013775108],"category_scores_gemma":[0.016095873,0.00040820808,0.0011644097,0.010831111,0.0015064477,0.0031427864,0.002995038,0.0012260952,0.00081514654],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036120322,0.0008100323,0.11263793,0.001165318,0.0004904523,0.0004081166,0.010600935,0.019601485,0.039288614,0.051126845,0.011711587,0.75179744],"study_design_scores_gemma":[0.000116401156,0.00017751189,0.081700906,0.0004188997,0.00024051924,0.0008982487,0.009844177,0.6747109,0.024335742,0.18079033,0.026475176,0.00029125958],"about_ca_topic_score_codex":0.008440545,"about_ca_topic_score_gemma":0.012823122,"teacher_disagreement_score":0.9808288,"about_ca_system_score_codex":0.0013263583,"about_ca_system_score_gemma":0.0023214687,"threshold_uncertainty_score":0.024897039},"labels":[],"label_agreement":null},{"id":"W2958192375","doi":"10.48550/arxiv.1307.8057","title":"Extracting Connected Concepts from Biomedical Texts using Fog Index","year":2013,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Readability; Categorization; Filter (signal processing); Computer science; Index (typography); Rank (graph theory); Artificial intelligence; Natural language processing; Value (mathematics); Measure (data warehouse); Harmonic mean; Mathematics; Data mining; Machine learning; Statistics; Combinatorics","score_opus":0.08242167131080438,"score_gpt":0.24212342747969892,"score_spread":0.15970175616889454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2958192375","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.43381494,0.008185676,0.502085,0.0015016163,0.00061179994,0.0014137827,0.025448289,0.012050107,0.014888966],"genre_scores_gemma":[0.5883323,0.001765431,0.38198918,0.00028297366,0.00072999747,0.0006189886,0.023476398,0.00031850635,0.002486203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986191,0.00014874383,0.00022767413,0.00037668456,0.00050668657,0.000121118435],"domain_scores_gemma":[0.99463516,0.002789902,0.0009382457,0.00040689943,0.000939549,0.00029026813],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014081569,0.0010395745,0.00096275314,0.025139527,0.0010734361,0.0019204052,0.00081014563,0.0012641756,0.0022517766],"category_scores_gemma":[0.0097932955,0.00027262518,0.0010709006,0.01028548,0.0006167076,0.0031278161,0.0014706756,0.0007954611,0.0013343813],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009939799,0.00039341152,0.060810033,0.0023472384,0.00048225655,0.0021762382,0.002300434,0.007222192,0.058320854,0.01083939,0.025132937,0.828981],"study_design_scores_gemma":[0.000263812,0.0010264501,0.25120312,0.000664859,0.0011059812,0.005398723,0.0032999797,0.43870243,0.071327716,0.12273536,0.10382196,0.00044957877],"about_ca_topic_score_codex":0.0025233354,"about_ca_topic_score_gemma":0.0031149772,"teacher_disagreement_score":0.025139527,"about_ca_system_score_codex":0.0007648082,"about_ca_system_score_gemma":0.0012155223,"threshold_uncertainty_score":0.007532954},"labels":[],"label_agreement":null},{"id":"W2963511785","doi":"","title":"Learning seasonal phytoplankton communities with topic models","year":2017,"lang":"en","type":"article","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Interpretability; Probabilistic logic; Statistical model; Computer science; Set (abstract data type); Regression analysis; Regression; Probability distribution; Machine learning; Artificial intelligence; Econometrics; Mathematics; Statistics","score_opus":0.07242095916689699,"score_gpt":0.19592797090672212,"score_spread":0.12350701173982513,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963511785","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09679724,0.00055562577,0.89844495,0.0007822356,0.00006274142,0.000069480535,0.00062871666,0.0011307977,0.0015283024],"genre_scores_gemma":[0.83654827,0.00064077304,0.1548452,0.00041548352,0.00029391248,0.0002843801,0.0027729128,0.00029720864,0.0039018954],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992093,0.00025824987,0.000038543596,0.00031230302,0.00009522743,0.000086400476],"domain_scores_gemma":[0.9965922,0.0025045876,0.00028113532,0.00025056518,0.0002399677,0.00013150374],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022196209,0.0008596224,0.0009976394,0.0018438515,0.0006853176,0.0013292942,0.002319452,0.0019036208,0.0018821644],"category_scores_gemma":[0.0081746,0.0010808927,0.0017126674,0.0016377611,0.00077754084,0.003206399,0.0016935604,0.0018476468,0.00079933787],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022744619,0.00017128048,0.015861105,0.00015950973,0.00025495415,0.00022379828,0.0006197195,0.8563728,0.0033873587,0.032703917,0.004725304,0.085292764],"study_design_scores_gemma":[0.000009481163,0.000007149948,0.0003702617,0.0000058313676,0.000008428427,0.00002120392,0.000017056509,0.9875568,0.0001698137,0.011458546,0.0003702368,0.000005237154],"about_ca_topic_score_codex":0.009123649,"about_ca_topic_score_gemma":0.014760597,"teacher_disagreement_score":0.009123649,"about_ca_system_score_codex":0.0011747042,"about_ca_system_score_gemma":0.0009052524,"threshold_uncertainty_score":0.01814109},"labels":[],"label_agreement":null},{"id":"W2963608472","doi":"","title":"Exploring Coverage and Distribution of Identifiers on the Scholarly Web","year":2015,"lang":"en","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Identifier; Computer science; World Wide Web; Publishing; Unique identifier; Distribution (mathematics); Information retrieval; Order (exchange); Web of science; Time lag; Data science; Lag; MEDLINE; Political science","score_opus":0.06411816845430199,"score_gpt":0.2535274075025609,"score_spread":0.18940923904825893,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2963608472","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9755089,0.004644384,0.0072971843,0.00068726164,0.00004293616,0.000028567629,0.005273961,0.00022862159,0.0062881853],"genre_scores_gemma":[0.9853196,0.0018600028,0.0041887336,0.00006218541,0.00011044287,0.000053462893,0.0069739833,0.00014626798,0.0012853249],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9884171,0.0032584958,0.0012611152,0.001696723,0.004701309,0.0006652154],"domain_scores_gemma":[0.8609135,0.10249174,0.014988535,0.005674799,0.014502781,0.0014285585],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.00867558,0.00033020184,0.0009683592,0.03169955,0.00075511483,0.0041804886,0.0008458162,0.0011476984,0.001888086],"category_scores_gemma":[0.101268,0.0002672493,0.0004726591,0.029936053,0.0010436211,0.0057944725,0.0031770861,0.0006610526,0.00081974536],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006368633,0.00012104896,0.7723013,0.0017002646,0.0006178056,0.0010271583,0.012692987,0.0038993766,0.008141862,0.0061675184,0.005927077,0.18676665],"study_design_scores_gemma":[0.000036991743,0.00019475377,0.8437638,0.0008876689,0.0005796501,0.0033133996,0.021205917,0.033258323,0.015436712,0.0151811065,0.06597198,0.00016964029],"about_ca_topic_score_codex":0.0025919753,"about_ca_topic_score_gemma":0.0015390386,"teacher_disagreement_score":0.9958195,"about_ca_system_score_codex":0.0006803332,"about_ca_system_score_gemma":0.00062571216,"threshold_uncertainty_score":0.04588139},"labels":[],"label_agreement":null},{"id":"W2964354311","doi":"10.1038/s41467-019-11026-x","title":"A machine-compiled database of genome-wide association studies","year":2019,"lang":"en","type":"article","venue":"Nature Communications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Defense Advanced Research Projects Agency; Office of Naval Research; Ant Financial Services Group; National Institutes of Health; National Science Foundation; VMware; Stanford Bio-X; Accenture; Gordon and Betty Moore Foundation; Okawa Foundation for Information and Telecommunications","keywords":"Computer science; Precision and recall; Information extraction; Information retrieval; Association (psychology); Genetic programming; Genome-wide association study; Knowledge base; Data science; Artificial intelligence; Genotype; Biology","score_opus":0.02464967991211376,"score_gpt":0.33187627681212045,"score_spread":0.3072265969000067,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2964354311","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013885497,0.005635801,0.025202759,0.0006671121,0.0002687898,0.00030327498,0.94032604,0.007515432,0.0061953585],"genre_scores_gemma":[0.026304102,0.003698921,0.05157773,0.00038538565,0.0001407371,0.0004991824,0.91542983,0.00059588853,0.0013682425],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9961104,0.0005437086,0.0013523767,0.0009934263,0.00084729056,0.00015288999],"domain_scores_gemma":[0.97764564,0.00996943,0.0039603966,0.003524678,0.003715226,0.0011847016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003972938,0.0012424132,0.0022584554,0.025173744,0.0011763318,0.003939948,0.0019865232,0.0013958054,0.012933213],"category_scores_gemma":[0.025844848,0.0009098025,0.0011399206,0.030562703,0.00059768016,0.002761672,0.0026718455,0.0018224941,0.0094585875],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002468604,0.0006306347,0.08354227,0.018542279,0.0028793386,0.0066929967,0.0016863999,0.008882949,0.04349178,0.021444382,0.4471323,0.36260605],"study_design_scores_gemma":[0.0005553981,0.00023257037,0.06933601,0.0017884413,0.0018604243,0.003126356,0.00048534203,0.0059957355,0.014718776,0.02344338,0.8781967,0.00026086418],"about_ca_topic_score_codex":0.0053294897,"about_ca_topic_score_gemma":0.01022546,"teacher_disagreement_score":0.025173744,"about_ca_system_score_codex":0.0011352451,"about_ca_system_score_gemma":0.006408875,"threshold_uncertainty_score":0.04326594},"labels":[],"label_agreement":null},{"id":"W2966783508","doi":"10.29173/cais1111","title":"Canada’s Health Data Repositories: Challenges of Organization, Discoverability and Access","year":2019,"lang":"en","type":"article","venue":"Proceedings of the Annual Conference of CAIS / Actes du congrès annuel de l ACSI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"University of Alberta","funders":"University of Alberta","keywords":"Discoverability; Metadata; Interoperability; Open data; Health data; World Wide Web; Computer science; Institutional repository; Linked data; Data science; Health care; Political science; Semantic Web","score_opus":0.031031206291311952,"score_gpt":0.2767233769537491,"score_spread":0.24569217066243715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2966783508","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1320338,0.027535696,0.14577572,0.53688514,0.0014722223,0.0019031924,0.041379496,0.0043337364,0.108681],"genre_scores_gemma":[0.5052388,0.019143108,0.38731888,0.020727595,0.00086019555,0.0010480711,0.04466121,0.0013244614,0.019677684],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9474993,0.013353874,0.004195708,0.0036970729,0.027329747,0.003924343],"domain_scores_gemma":[0.6914277,0.116318084,0.015531294,0.044689216,0.11467409,0.017359655],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.06802972,0.0007229167,0.0017031959,0.018076418,0.010868914,0.027764115,0.008288085,0.0026139354,0.0035374947],"category_scores_gemma":[0.16315886,0.0009475409,0.0017920493,0.03823815,0.009885949,0.019789396,0.01521294,0.003977518,0.0011732297],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004903297,0.00025035124,0.058633234,0.002890393,0.00060909754,0.0005096206,0.0142823355,0.005142072,0.002527049,0.250596,0.16488013,0.4991895],"study_design_scores_gemma":[0.00015013621,0.00012894429,0.0737556,0.004244389,0.00041424733,0.00058387883,0.028771313,0.013286017,0.003474736,0.15272133,0.7219124,0.0005571057],"about_ca_topic_score_codex":0.89959884,"about_ca_topic_score_gemma":0.91006434,"teacher_disagreement_score":0.97223586,"about_ca_system_score_codex":0.068363145,"about_ca_system_score_gemma":0.19214019,"threshold_uncertainty_score":0.49601167},"labels":[],"label_agreement":null},{"id":"W2967378256","doi":"","title":"Guides: Vancouver citation style (based on Citing Medicine): Thesis","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Citation; Computer science; History; World Wide Web; Archaeology","score_opus":0.042372180991692314,"score_gpt":0.28540133618699476,"score_spread":0.24302915519530244,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2967378256","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002502866,0.0007148307,0.022947175,0.0044613755,0.0029265324,0.00072084303,0.60321015,0.054852214,0.30766407],"genre_scores_gemma":[0.01116652,0.0017227543,0.0896627,0.0008453968,0.0011612125,0.000939906,0.35386935,0.038486388,0.5021457],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984806,0.00018543868,0.00029333148,0.00021653179,0.00074755383,0.00007666457],"domain_scores_gemma":[0.9789216,0.005776466,0.00091984094,0.0015534622,0.0115132695,0.0013154392],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0020637733,0.0012091392,0.0010296259,0.013749153,0.0018484524,0.0067996345,0.0016655037,0.0011815176,0.48761418],"category_scores_gemma":[0.02672586,0.00084329175,0.0005668152,0.02188756,0.00056660565,0.0033723577,0.0020589612,0.00134009,0.30271846],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020872933,0.000011748294,0.00017106492,0.0002728094,0.000004126376,0.000012783971,0.00007669914,0.000053900552,0.00012076484,0.0012708522,0.97421,0.02377445],"study_design_scores_gemma":[0.000022140115,0.0000066134276,0.0010109146,0.00015779091,0.000007307813,0.00004273211,0.00007628909,0.00039545924,0.0005881622,0.0018288816,0.99584717,0.000016546244],"about_ca_topic_score_codex":0.023336066,"about_ca_topic_score_gemma":0.07784216,"teacher_disagreement_score":0.99320036,"about_ca_system_score_codex":0.001762571,"about_ca_system_score_gemma":0.0057394556,"threshold_uncertainty_score":0.7308562},"labels":[],"label_agreement":null},{"id":"W2969402850","doi":"10.2196/13917","title":"Building a Semantic Health Data Warehouse in the Context of Clinical Trials: Development and Usability Study","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Data warehouse; Information retrieval; Ontology; Usability; Context (archaeology); SNOMED CT; Semantic interoperability; Health informatics; Semantic search; Terminology; Data science; World Wide Web; Semantic Web; Data mining; Medicine; Public health; Interoperability","score_opus":0.19201298191611957,"score_gpt":0.4852674325602752,"score_spread":0.2932544506441557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2969402850","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6448305,0.0013553934,0.3232208,0.003045433,0.00023840507,0.009026586,0.0053770393,0.007619728,0.0052860784],"genre_scores_gemma":[0.39771807,0.00060846657,0.5935347,0.00047308183,0.000037613612,0.0022739393,0.004502035,0.00043047057,0.00042153482],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.96620965,0.024359072,0.004315065,0.0013796225,0.0032643531,0.0004722398],"domain_scores_gemma":[0.8729504,0.09587046,0.0027184202,0.012251683,0.014536669,0.0016723712],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07669314,0.0009495639,0.0011479132,0.0038765357,0.0012037245,0.005719208,0.0021523046,0.0018446565,0.0012003406],"category_scores_gemma":[0.103222445,0.00096607296,0.0021771933,0.0033088543,0.001206083,0.0069568777,0.0042279786,0.0017106395,0.00048499517],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0055774436,0.0077935704,0.10162115,0.014366191,0.002151206,0.004929598,0.09638667,0.023637988,0.06061458,0.02350871,0.019672368,0.6397405],"study_design_scores_gemma":[0.003916841,0.011142895,0.11583387,0.009273475,0.0038405573,0.008129314,0.08653369,0.35254368,0.0995213,0.03440966,0.27333453,0.0015201566],"about_ca_topic_score_codex":0.0023398197,"about_ca_topic_score_gemma":0.0021539961,"teacher_disagreement_score":0.07669314,"about_ca_system_score_codex":0.0016404744,"about_ca_system_score_gemma":0.004966798,"threshold_uncertainty_score":0.4055969},"labels":[],"label_agreement":null},{"id":"W2969446777","doi":"","title":"Guides: Vancouver reference style (based on Citing Medicine): ABS and AIHW publications","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); History; Archaeology","score_opus":0.05819877255824472,"score_gpt":0.30234461117448663,"score_spread":0.24414583861624192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2969446777","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008000665,0.0013352372,0.01074513,0.0072369347,0.0063305353,0.0006748734,0.3962469,0.02023335,0.55639696],"genre_scores_gemma":[0.0024741914,0.0026677472,0.0248258,0.0012022534,0.0013138902,0.0006250914,0.2484525,0.016448427,0.70199],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965322,0.0003204154,0.0006726922,0.00039903307,0.0018894874,0.00018609822],"domain_scores_gemma":[0.95895296,0.0064960723,0.0014688846,0.0023747792,0.028409299,0.0022981004],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003392602,0.0019269867,0.0021668877,0.028504986,0.0032076337,0.011166921,0.0029462145,0.00230528,0.57944566],"category_scores_gemma":[0.036346372,0.0012958649,0.0008921573,0.05090727,0.0010665266,0.0066529736,0.00331035,0.0024654113,0.5581938],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007154599,0.000006145931,0.00004051092,0.00015180842,0.0000016982889,0.000008520136,0.000034268596,0.00002287466,0.000041112587,0.0006359647,0.9880965,0.010953492],"study_design_scores_gemma":[0.000007428979,0.0000033982485,0.0004912125,0.00023114427,0.0000043591367,0.00003077594,0.00007794861,0.000093546085,0.0001442085,0.00091178884,0.99799216,0.000012069297],"about_ca_topic_score_codex":0.070700526,"about_ca_topic_score_gemma":0.16102466,"teacher_disagreement_score":0.57944566,"about_ca_system_score_codex":0.0033783976,"about_ca_system_score_gemma":0.01041705,"threshold_uncertainty_score":0.5998697},"labels":[],"label_agreement":null},{"id":"W2969593549","doi":"","title":"Guides: Vancouver citation style (based on Citing Medicine): Social media","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Social media; Citation; Media studies; Internet privacy; Sociology; Computer science; World Wide Web; Art; Visual arts","score_opus":0.05055945812572663,"score_gpt":0.29263787452476203,"score_spread":0.2420784163990354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2969593549","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068827067,0.0010000926,0.020890247,0.0020198459,0.00060936215,0.0005707202,0.71340483,0.07419911,0.18042308],"genre_scores_gemma":[0.020411735,0.0019298128,0.08592063,0.0004450187,0.0003949773,0.00079308957,0.52634454,0.025918005,0.3378422],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989311,0.00011574835,0.00016438578,0.00015033875,0.0005774528,0.000061016148],"domain_scores_gemma":[0.99178934,0.0022817492,0.00049997034,0.00066900806,0.004090688,0.00066923077],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0013160995,0.001321712,0.00068592886,0.012863015,0.0012363002,0.0044101216,0.0012871638,0.0009891541,0.18342125],"category_scores_gemma":[0.01185345,0.000782903,0.00050592533,0.019647043,0.0004341599,0.002423601,0.0016345693,0.00088868005,0.101365775],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055456978,0.00003726452,0.0009212663,0.00055539014,0.000016186714,0.000037283895,0.00017697913,0.00022772774,0.0003048837,0.0023262494,0.9250526,0.0702888],"study_design_scores_gemma":[0.00004627307,0.000017124188,0.0041432977,0.00018645948,0.000024970832,0.00009357246,0.00013207502,0.0017552779,0.0016022229,0.0028176075,0.9891427,0.00003846202],"about_ca_topic_score_codex":0.07212339,"about_ca_topic_score_gemma":0.1778682,"teacher_disagreement_score":0.99558985,"about_ca_system_score_codex":0.0016266739,"about_ca_system_score_gemma":0.0045104544,"threshold_uncertainty_score":0.61360526},"labels":[],"label_agreement":null},{"id":"W2970205171","doi":"","title":"TAC SRIE 2018: Extracting Systematic Review Information with MedaCy.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Information retrieval","score_opus":0.006880278853483287,"score_gpt":0.2608429318200521,"score_spread":0.2539626529665688,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970205171","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009196787,0.02906422,0.10601774,0.004465314,0.0010392166,0.006065979,0.8111305,0.0223273,0.010692906],"genre_scores_gemma":[0.04111003,0.009383436,0.50150275,0.0017127118,0.0005181598,0.021724327,0.41821876,0.0020997133,0.0037300675],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9859995,0.0052578785,0.0047663376,0.0019043336,0.0018378962,0.0002340276],"domain_scores_gemma":[0.90262336,0.07482126,0.007640764,0.007882181,0.0056633535,0.0013690842],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018306829,0.0017999241,0.0038079964,0.033274915,0.0019425747,0.0054012155,0.002093149,0.001852697,0.030343521],"category_scores_gemma":[0.16238123,0.0012703884,0.005623299,0.021832222,0.0012688296,0.004732803,0.0069323527,0.002003831,0.008231585],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016785274,0.00018150274,0.017923098,0.19353592,0.010645001,0.0011694762,0.0040544337,0.0019354349,0.007172657,0.025227992,0.40312895,0.33334696],"study_design_scores_gemma":[0.0010759128,0.00036568652,0.016910193,0.023600083,0.010794891,0.0012668747,0.0012298343,0.0073116384,0.0051579536,0.06142821,0.8705356,0.0003231416],"about_ca_topic_score_codex":0.0032452417,"about_ca_topic_score_gemma":0.014319386,"teacher_disagreement_score":0.98169315,"about_ca_system_score_codex":0.0016368998,"about_ca_system_score_gemma":0.012426553,"threshold_uncertainty_score":0.10150921},"labels":[],"label_agreement":null},{"id":"W2970391280","doi":"","title":"A Pragmatic Approach to Information Extraction for Systematic Review.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Information extraction; Extraction (chemistry); Information retrieval; Data science; Chemistry; Chromatography","score_opus":0.008141321295185574,"score_gpt":0.2920434110316978,"score_spread":0.2839020897365122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970391280","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022215392,0.025697967,0.86896235,0.027162733,0.00219248,0.04106911,0.019041875,0.002713784,0.010938156],"genre_scores_gemma":[0.011355428,0.0024781835,0.94939625,0.0018580429,0.00029036475,0.031349424,0.0025286847,0.000115066614,0.00062857737],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.5578824,0.33354053,0.07063773,0.009237959,0.027495267,0.0012060995],"domain_scores_gemma":[0.3446125,0.5611082,0.022115277,0.028220074,0.041182615,0.0027613295],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2672213,0.0031021172,0.007315247,0.03150138,0.0051938556,0.013374627,0.0058563035,0.005472505,0.01578308],"category_scores_gemma":[0.6007952,0.0042464635,0.00930744,0.02649818,0.0045737233,0.010090686,0.015463044,0.0064518116,0.0039619934],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001769381,0.0001754455,0.001869344,0.22646014,0.011845661,0.000990778,0.0126203,0.0029373972,0.003824627,0.17180665,0.0962806,0.46941966],"study_design_scores_gemma":[0.0021552746,0.00031003528,0.0031300122,0.07976191,0.01282112,0.0011253079,0.0038253102,0.013839519,0.0039278553,0.5543725,0.32399604,0.0007352584],"about_ca_topic_score_codex":0.0045177587,"about_ca_topic_score_gemma":0.013678184,"teacher_disagreement_score":0.73277867,"about_ca_system_score_codex":0.007579953,"about_ca_system_score_gemma":0.044176575,"threshold_uncertainty_score":0.9036466},"labels":[],"label_agreement":null},{"id":"W2970640551","doi":"","title":"Overview of the TAC 2018 Drug-Drug Interaction Extraction from Drug Labels Track.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug; Drug-drug interaction; Computer science; Track (disk drive); Extraction (chemistry); Pharmacology; Medicine; Chemistry; Chromatography","score_opus":0.015625761155174986,"score_gpt":0.3036536047877693,"score_spread":0.2880278436325943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970640551","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013224577,0.017564762,0.46054456,0.0038313833,0.0015252738,0.0017557837,0.3193325,0.14135273,0.04086839],"genre_scores_gemma":[0.029473677,0.0045103133,0.34282726,0.0011011545,0.00033705766,0.0011461437,0.60643786,0.0027386346,0.011427859],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975242,0.00038396943,0.0003195275,0.0005616362,0.0010129048,0.0001977742],"domain_scores_gemma":[0.99569726,0.0010251023,0.0002752878,0.0011081515,0.0015946133,0.00029968028],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028991355,0.0014519915,0.0015026589,0.011305724,0.0014377286,0.0035265575,0.0024065583,0.001462394,0.011915909],"category_scores_gemma":[0.009421013,0.00065411936,0.001978549,0.007841637,0.00044921867,0.0036200255,0.002535718,0.0018397868,0.014833181],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046961824,0.0002829407,0.004891409,0.0022963372,0.00038399152,0.0003095187,0.00018011054,0.0048410404,0.009893735,0.011179468,0.5391916,0.4260802],"study_design_scores_gemma":[0.0001413293,0.00026543307,0.00887056,0.0006898855,0.00040770654,0.00082279457,0.00018084639,0.081975214,0.015198435,0.024261681,0.86706215,0.00012398786],"about_ca_topic_score_codex":0.021320613,"about_ca_topic_score_gemma":0.031636305,"teacher_disagreement_score":0.021320613,"about_ca_system_score_codex":0.0016325215,"about_ca_system_score_gemma":0.0056294748,"threshold_uncertainty_score":0.04239303},"labels":[],"label_agreement":null},{"id":"W2970660638","doi":"10.2196/12575","title":"Extracting Clinical Features From Dictated Ambulatory Consult Notes Using a Commercially Available Natural Language Processing Tool: Pilot, Retrospective, Cross-Sectional Validation Study","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"St. Michael's Hospital; Hamilton Health Sciences; Public Health Ontario; University of Toronto","funders":"","keywords":"Natural language processing; Artificial intelligence; Computer science; Abstraction; Information retrieval; Medicine; Health records; Machine learning; Health care","score_opus":0.04302415869666514,"score_gpt":0.38672845985186016,"score_spread":0.343704301155195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2970660638","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9958682,0.0001161133,0.002622292,0.00003277335,0.000009097212,0.00067606353,0.0003958689,0.000023553026,0.00025609235],"genre_scores_gemma":[0.99113363,0.0001175868,0.0065099103,0.00014593343,0.000019187886,0.0007825302,0.0011586832,0.000016124937,0.00011652707],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9918795,0.0033812455,0.0013058123,0.0013117128,0.0018554255,0.00026624132],"domain_scores_gemma":[0.9494976,0.026166746,0.007531772,0.005413356,0.010365802,0.0010247001],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0138401855,0.00050700887,0.00044242173,0.001227253,0.00069226057,0.00076360116,0.00062437414,0.00061016536,0.00073807046],"category_scores_gemma":[0.043595884,0.0004454761,0.00049367204,0.0008026937,0.0010151167,0.0008909167,0.0008915569,0.00064247265,0.0004531246],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069076143,0.0012314231,0.9793729,0.00017261133,0.000099734134,0.00029062195,0.0026692825,0.00018249342,0.0025239394,0.000048136837,0.00033847685,0.0123797115],"study_design_scores_gemma":[0.00015574029,0.005652845,0.98507684,0.0000914665,0.00011909123,0.0014727158,0.0015776054,0.0017345264,0.0029013213,0.000053186122,0.0011353068,0.000029445257],"about_ca_topic_score_codex":0.0023183024,"about_ca_topic_score_gemma":0.003526856,"teacher_disagreement_score":0.0138401855,"about_ca_system_score_codex":0.00063379103,"about_ca_system_score_gemma":0.0015032028,"threshold_uncertainty_score":0.07319474},"labels":[],"label_agreement":null},{"id":"W2971077097","doi":"","title":"IBM Research System at TAC 2018: Deep Learning architectures for Drug-Drug Interaction extraction from Structured Product Labels.","year":2018,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"IBM; Computer science; Extraction (chemistry); Drug; Artificial intelligence; Product (mathematics); Chemistry; Pharmacology; Chromatography; Mathematics; Medicine; Nanotechnology; Materials science","score_opus":0.016578716248759333,"score_gpt":0.3265698188233774,"score_spread":0.30999110257461804,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971077097","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07317344,0.011365595,0.36484373,0.004335729,0.0012835099,0.0006768039,0.14030623,0.38064635,0.023368541],"genre_scores_gemma":[0.2193848,0.0033496197,0.50359005,0.0014329077,0.0003260639,0.0014960753,0.23087235,0.006776227,0.03277191],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995454,0.000109145534,0.000032149008,0.00014264126,0.00011970773,0.000050904564],"domain_scores_gemma":[0.9991898,0.00029288855,0.000059003483,0.00017888864,0.00020768335,0.00007178687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011952758,0.0016746633,0.0009842367,0.0015814984,0.00057154935,0.0013516576,0.0022356375,0.0013595873,0.018819802],"category_scores_gemma":[0.004469621,0.00084419333,0.0010299368,0.0020294008,0.00030025374,0.0022651174,0.0014674971,0.0022226062,0.011918845],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014205199,0.000599029,0.0027584075,0.0010631905,0.0005422271,0.00022028411,0.00017322696,0.03629625,0.0084297815,0.0073944223,0.5642994,0.37680334],"study_design_scores_gemma":[0.00087790925,0.0005402424,0.0018460271,0.00018832421,0.0002743495,0.00016776694,0.000093467956,0.79994684,0.01918204,0.036307625,0.14048815,0.00008731119],"about_ca_topic_score_codex":0.017762877,"about_ca_topic_score_gemma":0.02617709,"teacher_disagreement_score":0.018819802,"about_ca_system_score_codex":0.0013467644,"about_ca_system_score_gemma":0.0031211448,"threshold_uncertainty_score":0.06295854},"labels":[],"label_agreement":null},{"id":"W2971337214","doi":"10.3747/co.26.5535","title":"Impact of the Knowledge Translation Research Network’s Grants Program in Cancer Knowledge Translation","year":2019,"lang":"en","type":"editorial","venue":"Current Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University Health Network; Ontario Institute for Cancer Research","funders":"","keywords":"Knowledge translation; Translation (biology); Medicine; Cancer; Translational research; Knowledge management; Computer science; Pathology; Genetics; Biology; Messenger RNA","score_opus":0.306960057558224,"score_gpt":0.5531828967510783,"score_spread":0.24622283919285437,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2971337214","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00001634487,0.0045697726,0.00013230831,0.1047779,0.8892832,0.000015784955,0.00008187085,0.00003785339,0.0010850512],"genre_scores_gemma":[0.00045181715,0.006510981,0.00026126014,0.07141129,0.91300106,0.000045278444,0.00006745864,0.000073603944,0.008177366],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98417974,0.0039432053,0.0022265029,0.0014268941,0.007433318,0.00079033733],"domain_scores_gemma":[0.8714338,0.07304536,0.0037621064,0.0019015997,0.040632915,0.009224293],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.028785333,0.0026246482,0.0031539642,0.0058795135,0.004505453,0.0126732215,0.004557576,0.021102622,0.015133301],"category_scores_gemma":[0.09228831,0.0012033719,0.0034151198,0.0031247118,0.0046123583,0.008152813,0.0034987738,0.032244757,0.009529922],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017562765,0.000003882761,0.000010931979,0.00013015178,0.000010976482,0.00003616194,0.000013313819,0.000013292036,0.000012267792,0.0005046787,0.9965522,0.0026946182],"study_design_scores_gemma":[0.00006976726,0.00001850347,0.0001580225,0.0010645965,0.00006642994,0.00013401786,0.00004696275,0.0001430048,0.00007579161,0.0019036282,0.99629164,0.000027690508],"about_ca_topic_score_codex":0.006454007,"about_ca_topic_score_gemma":0.015438347,"teacher_disagreement_score":0.99161565,"about_ca_system_score_codex":0.008384345,"about_ca_system_score_gemma":0.013207568,"threshold_uncertainty_score":0.15223318},"labels":[],"label_agreement":null},{"id":"W2972887977","doi":"10.4220/sykepleienf.2019.78413","title":"Kartlegging av prosedyrer for oppdekking av instrumentbord ved kirurgiske inngrep ","year":2019,"lang":"no","type":"article","venue":"Sykepleien Forskning","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Washington; University of Toronto; Massachusetts General Hospital; University of the Philippines; Imperial College Healthcare NHS Trust; Imperial College London; Brigham and Women's Hospital","keywords":"Medicine","score_opus":0.016473270241698625,"score_gpt":0.2776125714705734,"score_spread":0.26113930122887474,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2972887977","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039272983,0.011309924,0.017569412,0.047978304,0.015451123,0.0008896928,0.003818583,0.0028496706,0.8608602],"genre_scores_gemma":[0.058475684,0.007077033,0.014095115,0.0048830085,0.001311784,0.0004952465,0.002325726,0.0018069441,0.9095294],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9956929,0.0006078186,0.00024841607,0.00061599474,0.0020013466,0.0008335306],"domain_scores_gemma":[0.9939009,0.0008301037,0.0004156708,0.00035917843,0.0014844973,0.0030096697],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003530734,0.0012876262,0.0009253584,0.0017281214,0.0035183628,0.0066121286,0.001920702,0.0029055371,0.3221664],"category_scores_gemma":[0.007599422,0.0008326625,0.0014244472,0.00078141486,0.0018344667,0.0037359293,0.009942139,0.0052846917,0.15452594],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000785387,0.0005925089,0.005320404,0.0017838958,0.000049635146,0.00097371725,0.006245972,0.00080536795,0.0077534365,0.012148914,0.5612687,0.40227208],"study_design_scores_gemma":[0.000030739935,0.00021921363,0.0061269654,0.0008396765,0.000018676834,0.0004993223,0.0033633474,0.00025383773,0.0017252752,0.0028014132,0.9840581,0.00006356949],"about_ca_topic_score_codex":0.0062197773,"about_ca_topic_score_gemma":0.019228484,"teacher_disagreement_score":0.3221664,"about_ca_system_score_codex":0.0037765845,"about_ca_system_score_gemma":0.009782393,"threshold_uncertainty_score":0.96684736},"labels":[],"label_agreement":null},{"id":"W2973398987","doi":"10.1186/s13326-019-0207-3","title":"Automated SNOMED CT concept and attribute relationship detection through a web-based implementation of cTAKES","year":2019,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"European Regional Development Fund","keywords":"SNOMED CT; Computer science; Information retrieval; World Wide Web; Data science; Data mining; Terminology","score_opus":0.016632121513990142,"score_gpt":0.309345508240613,"score_spread":0.2927133867266229,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973398987","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18572675,0.00045037057,0.51452565,0.00027708264,0.00014205652,0.0012336025,0.007939271,0.28695637,0.002748864],"genre_scores_gemma":[0.2848354,0.00020164534,0.6951294,0.0003185985,0.000037853253,0.00072668825,0.0104057565,0.0035733287,0.004771434],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99857104,0.0002989139,0.00020970574,0.00060374255,0.00025850662,0.000058020752],"domain_scores_gemma":[0.9889269,0.006915999,0.0007135693,0.001262801,0.0018034822,0.0003772075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025564546,0.0009864692,0.0008329989,0.002694433,0.00045452098,0.0014325168,0.0017534673,0.0010427113,0.0065869386],"category_scores_gemma":[0.010662058,0.0008025224,0.0007882185,0.0013113525,0.00049120816,0.0019919977,0.0015464415,0.0007523581,0.0022187044],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0074414895,0.0014375536,0.033991836,0.0032555764,0.0009140225,0.0027756235,0.002180498,0.02636929,0.15646355,0.0048964173,0.045684308,0.71458983],"study_design_scores_gemma":[0.0009676734,0.0013050939,0.033877537,0.0003567035,0.00041948276,0.0040999376,0.00063013966,0.6159472,0.2622027,0.0038806621,0.075873256,0.00043964476],"about_ca_topic_score_codex":0.0038778528,"about_ca_topic_score_gemma":0.005348193,"teacher_disagreement_score":0.0065869386,"about_ca_system_score_codex":0.0009807718,"about_ca_system_score_gemma":0.0016084472,"threshold_uncertainty_score":0.02203548},"labels":[],"label_agreement":null},{"id":"W2973539856","doi":"10.1038/s41538-019-0048-6","title":"Global agricultural concept space: lightweight semantics for pragmatic interoperability","year":2019,"lang":"en","type":"article","venue":"npj Science of Food","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Agriculture and Agri-Food Canada; Consortium of International Agricultural Research Centers; Institut National de la Recherche Agronomique; Ministry of Agriculture of the People's Republic of China; Department for International Development","keywords":"Computer science; Identifier; Metadata; Thesaurus; Information retrieval; Interoperability; Semantics (computer science); Space (punctuation); Ontology; World Wide Web; Semantic Web; Artificial intelligence; Programming language","score_opus":0.008770027599883708,"score_gpt":0.2680695043544615,"score_spread":0.25929947675457776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2973539856","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006232031,0.0004123659,0.9820393,0.0018224218,0.00013433037,0.00023794477,0.0014726814,0.0013695322,0.0062792953],"genre_scores_gemma":[0.13716726,0.00059184164,0.8542041,0.00054576143,0.0001800769,0.0006473726,0.0036813542,0.0005086425,0.0024736607],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9920352,0.003710494,0.0014259666,0.0010733429,0.0014081176,0.00034682712],"domain_scores_gemma":[0.99066657,0.0040114033,0.0007342175,0.0030155357,0.0011629662,0.00040919488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0094282525,0.0009833369,0.00105172,0.0061312104,0.002382898,0.00834627,0.0020594236,0.001897851,0.0040809666],"category_scores_gemma":[0.016743612,0.0008154662,0.0028465334,0.0051805116,0.005785093,0.019799368,0.008388846,0.00297724,0.0013247897],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004337425,0.000021848873,0.00036062748,0.0002050738,0.000034611614,0.00011718778,0.0018068021,0.0015107875,0.0011796838,0.961214,0.0030260102,0.030479923],"study_design_scores_gemma":[0.000017234044,0.000025710999,0.00030941478,0.00017389741,0.00004564175,0.00019369174,0.00094987015,0.012938435,0.0018253635,0.8923193,0.091163345,0.000038130493],"about_ca_topic_score_codex":0.00477709,"about_ca_topic_score_gemma":0.004185359,"teacher_disagreement_score":0.0094282525,"about_ca_system_score_codex":0.002481769,"about_ca_system_score_gemma":0.0037702196,"threshold_uncertainty_score":0.049861908},"labels":[],"label_agreement":null},{"id":"W2974711343","doi":"","title":"LibGuides: MHIKNET: Order Documents","year":2019,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Order (exchange); Resource (disambiguation); World Wide Web; Computer science; History; Business","score_opus":0.016186417266834736,"score_gpt":0.28786006795387814,"score_spread":0.2716736506870434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2974711343","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00090907095,0.00046754794,0.006498966,0.002104921,0.00047937923,0.00022928628,0.2689918,0.15583219,0.56448686],"genre_scores_gemma":[0.00855622,0.0010748954,0.011123607,0.0010903043,0.00041644427,0.00039535965,0.36540964,0.084108815,0.5278247],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99895585,0.0001246944,0.00010627157,0.00014856826,0.00052287075,0.00014172941],"domain_scores_gemma":[0.99519545,0.0009981229,0.0002669755,0.0010062668,0.0018022712,0.00073086546],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0013329104,0.001943939,0.001116623,0.0071480256,0.0020618867,0.008347734,0.0031455962,0.0022407444,0.7237346],"category_scores_gemma":[0.0112117175,0.0015108595,0.0006730016,0.009454684,0.00089043507,0.008726349,0.0050479714,0.0013987541,0.73441607],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003789509,0.000015578491,0.0000900451,0.00017092071,0.0000026098658,0.00002621284,0.00007109417,0.000039055714,0.00015171486,0.0024195556,0.9716404,0.025334848],"study_design_scores_gemma":[0.000014745202,0.000004573979,0.00016696786,0.000055121895,0.000001805198,0.000035384954,0.000062417756,0.00014073761,0.00040551223,0.0014931725,0.9976063,0.0000131951065],"about_ca_topic_score_codex":0.021106942,"about_ca_topic_score_gemma":0.02363071,"teacher_disagreement_score":0.7237346,"about_ca_system_score_codex":0.003631365,"about_ca_system_score_gemma":0.0033801815,"threshold_uncertainty_score":0.394059},"labels":[],"label_agreement":null},{"id":"W2975135115","doi":"10.1016/j.jbi.2019.103292","title":"Wikidata: A large-scale collaborative ontological medical database","year":2019,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":61,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"University of Virginia","keywords":"Computer science; Interoperability; Multidisciplinary approach; Data science; Resource (disambiguation); Knowledge base; World Wide Web; Scale (ratio); Ontology; Information retrieval; Database","score_opus":0.008969678556482182,"score_gpt":0.28411783339198315,"score_spread":0.275148154835501,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2975135115","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032527216,0.0031954118,0.61992407,0.0033822402,0.0011576134,0.0026163608,0.15307353,0.16399997,0.02012358],"genre_scores_gemma":[0.1172073,0.0027936515,0.54983914,0.0021433767,0.00030701942,0.0022503212,0.30332494,0.010367684,0.011766582],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968464,0.0005327976,0.0007267057,0.000749588,0.0010112236,0.00013322201],"domain_scores_gemma":[0.9884263,0.0043247105,0.00079782907,0.0037660801,0.0011501256,0.0015349999],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004857904,0.0015396556,0.002403722,0.008214477,0.0018421751,0.0052117286,0.00423523,0.0015418226,0.010442144],"category_scores_gemma":[0.019924104,0.0011445677,0.002083162,0.005952488,0.00067989377,0.006938737,0.010799151,0.0024201665,0.0060156244],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035025033,0.0016001988,0.017801968,0.0053633926,0.004216519,0.00282889,0.002956652,0.012601368,0.027378233,0.05439537,0.34086737,0.5264875],"study_design_scores_gemma":[0.0012308538,0.00047297907,0.013827946,0.00089817814,0.002292451,0.0031513388,0.0018750407,0.13559912,0.025572727,0.10357023,0.7108432,0.0006659183],"about_ca_topic_score_codex":0.0076217256,"about_ca_topic_score_gemma":0.01375711,"teacher_disagreement_score":0.010442144,"about_ca_system_score_codex":0.0011114966,"about_ca_system_score_gemma":0.005958425,"threshold_uncertainty_score":0.034932435},"labels":[],"label_agreement":null},{"id":"W2975460762","doi":"10.1177/1073110519876166","title":"Cochrane's Linked Data Project: How it Can Advance our Understanding of Surrogate Endpoints","year":2019,"lang":"en","type":"article","venue":"The Journal of Law Medicine & Ethics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cochrane","funders":"","keywords":"Surrogate endpoint; Data science; Computer science; Real world evidence; Medicine; Internal medicine","score_opus":0.23744184286441605,"score_gpt":0.42571706817732763,"score_spread":0.18827522531291158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2975460762","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036780227,0.057665065,0.32757285,0.20321207,0.0064703203,0.007878451,0.27247688,0.02394656,0.09709977],"genre_scores_gemma":[0.020822357,0.039129704,0.77339256,0.020879386,0.0020117536,0.013150699,0.1111625,0.011029632,0.008421397],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.8474099,0.099447526,0.025215484,0.005748417,0.020369474,0.0018091598],"domain_scores_gemma":[0.3265794,0.5443191,0.019570492,0.06401379,0.036222067,0.009295179],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.17796804,0.0021423358,0.0038988744,0.042631045,0.0035247128,0.020339346,0.0058627455,0.011068036,0.039983246],"category_scores_gemma":[0.57185996,0.0028620919,0.007646438,0.048012953,0.0048952885,0.016031794,0.019149553,0.009283784,0.008996368],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012699887,0.00021490022,0.0031493043,0.020502074,0.0020600527,0.00036169623,0.0024009775,0.0045998446,0.000484296,0.28295553,0.45871913,0.22328217],"study_design_scores_gemma":[0.00079157535,0.00008959042,0.0013825543,0.01948227,0.0010935623,0.00021160282,0.0003664175,0.0026190998,0.00088078936,0.12332084,0.8494946,0.00026723303],"about_ca_topic_score_codex":0.04840533,"about_ca_topic_score_gemma":0.03013691,"teacher_disagreement_score":0.822032,"about_ca_system_score_codex":0.008840867,"about_ca_system_score_gemma":0.06800215,"threshold_uncertainty_score":0.9411962},"labels":[],"label_agreement":null},{"id":"W2976978765","doi":"10.2196/13430","title":"Impact of Automatic Query Generation and Quality Recognition Using Deep Learning to Curate Evidence From Biomedical Literature: Empirical Study","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Quality (philosophy); Artificial intelligence; Natural language processing; Deep learning; Information retrieval; Data science","score_opus":0.08052265100987234,"score_gpt":0.41628811569680113,"score_spread":0.3357654646869288,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2976978765","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9280614,0.017002812,0.04017336,0.0034492726,0.0002834745,0.0018317858,0.0034835637,0.0012770065,0.004437371],"genre_scores_gemma":[0.9574539,0.001371325,0.0356378,0.0005869043,0.00014665961,0.0004233742,0.0037047993,0.00007062945,0.00060456124],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9603549,0.02415763,0.003907668,0.003541165,0.0072997888,0.0007388851],"domain_scores_gemma":[0.7120537,0.25173137,0.0118318945,0.009560794,0.012394422,0.0024277633],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04049787,0.001023546,0.0012933174,0.0042389757,0.0010114763,0.0025254318,0.0024056195,0.0017985522,0.0017729738],"category_scores_gemma":[0.18676966,0.00042192606,0.0018148533,0.0030266608,0.0012827012,0.004412067,0.0021588595,0.0022169256,0.0006603937],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005793207,0.0066870665,0.43340966,0.007171227,0.0024332716,0.0005662659,0.0016715244,0.015654044,0.0038454172,0.001443673,0.01544935,0.50587523],"study_design_scores_gemma":[0.0016227423,0.005606609,0.24556068,0.0023412504,0.0058337855,0.0022714715,0.0041113095,0.68929374,0.015135441,0.008824274,0.01906928,0.00032949387],"about_ca_topic_score_codex":0.008385329,"about_ca_topic_score_gemma":0.009304315,"teacher_disagreement_score":0.9595021,"about_ca_system_score_codex":0.002132496,"about_ca_system_score_gemma":0.004240357,"threshold_uncertainty_score":0.21417576},"labels":[],"label_agreement":null},{"id":"W2977770694","doi":"10.2210/pdb3ghj/pdb","title":"Crystal structure from the mobile metagenome of Halifax Harbour Sewage Outfall: Integron Cassette Protein HFX_CASS4","year":2009,"lang":"en","type":"paratext","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Harbour; Outfall; Integron; Sewage; Environmental science; Geology; Environmental engineering; Computer science; Paleontology","score_opus":0.012143654946168862,"score_gpt":0.26240654408007813,"score_spread":0.25026288913390926,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2977770694","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6532847,0.004237249,0.0087414645,0.0052832603,0.00062488957,0.0002111048,0.29940394,0.0064454563,0.021767868],"genre_scores_gemma":[0.2781282,0.0054511465,0.019216213,0.00047267313,0.000069267044,0.00026774226,0.67126906,0.001457088,0.02366862],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9998528,0.000007723062,0.000012862891,0.000033051667,0.00006850777,0.000024985087],"domain_scores_gemma":[0.999742,0.000101636026,0.00003541012,0.000016739914,0.00005846655,0.000045688652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00026504425,0.0016213183,0.0012164802,0.0008064255,0.00093385816,0.0016218207,0.0014195491,0.001297872,0.017128387],"category_scores_gemma":[0.0007996495,0.0006295656,0.00069470424,0.0019591071,0.00027452168,0.00076094177,0.00059701066,0.0017160991,0.0033635425],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005022299,0.0009759309,0.020211024,0.0058804406,0.00092427345,0.0082374215,0.0029602018,0.010444534,0.5208707,0.006816478,0.35850957,0.05914717],"study_design_scores_gemma":[0.002373292,0.0010558079,0.13669878,0.00072939176,0.0013500082,0.0040035425,0.0032576148,0.07323378,0.2560887,0.004185833,0.5165444,0.0004788968],"about_ca_topic_score_codex":0.03924628,"about_ca_topic_score_gemma":0.04894234,"teacher_disagreement_score":0.03924628,"about_ca_system_score_codex":0.0023244317,"about_ca_system_score_gemma":0.0019562375,"threshold_uncertainty_score":0.07803571},"labels":[],"label_agreement":null},{"id":"W2979820532","doi":"10.1093/bioinformatics/btz744","title":"Soft windowing application to improve analysis of high-throughput phenotyping data","year":2019,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lunenfeld-Tanenbaum Research Institute; Mount Sinai Hospital; University of Manitoba; Toronto Centre for Phenogenomics; Hospital for Sick Children","funders":"Deutsches Zentrum für Diabetesforschung; National Human Genome Research Institute; National Institutes of Health; Agence Nationale de la Recherche; European Molecular Biology Laboratory","keywords":"Computer science; Throughput; Software; Data mining; Computational biology; Programming language; Biology; Operating system","score_opus":0.015526743440320546,"score_gpt":0.27634203815419134,"score_spread":0.26081529471387077,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2979820532","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017747931,0.0004322646,0.9293454,0.0004084,0.00020924232,0.00019253597,0.0031763199,0.047881708,0.00060617103],"genre_scores_gemma":[0.05533906,0.00023793723,0.929635,0.00028616446,0.00009937912,0.00063481525,0.004413248,0.008281447,0.0010728928],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953477,0.0018352509,0.00040307152,0.0011212268,0.0010521631,0.00024061523],"domain_scores_gemma":[0.9762981,0.017035669,0.0014515428,0.0028860716,0.0017783844,0.00055034197],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012584967,0.0019518046,0.0016347427,0.0026007623,0.0009244692,0.0025657082,0.0025919233,0.0014126457,0.011828153],"category_scores_gemma":[0.047586642,0.0013848639,0.0026003874,0.002650555,0.0008741265,0.0022981807,0.003200082,0.0031468198,0.0041145436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028864383,0.00046975267,0.034217622,0.0029385397,0.0022331153,0.0015015965,0.0017879733,0.09751954,0.1640367,0.018211104,0.10711815,0.5670795],"study_design_scores_gemma":[0.000398484,0.00036859262,0.018825663,0.0002123629,0.0003759962,0.0008149876,0.00019352429,0.79595256,0.09165774,0.03743505,0.05344804,0.00031703085],"about_ca_topic_score_codex":0.002725528,"about_ca_topic_score_gemma":0.0034733897,"teacher_disagreement_score":0.012584967,"about_ca_system_score_codex":0.0007490545,"about_ca_system_score_gemma":0.00216099,"threshold_uncertainty_score":0.06655645},"labels":[],"label_agreement":null},{"id":"W2980234588","doi":"10.5858/arpa.2019-0276-oa","title":"A Survey of LOINC Code Selection Practices Among Participants of the College of American Pathologists Coagulation (CGL) and Cardiac Markers (CRT) Proficiency Testing Programs","year":2019,"lang":"en","type":"article","venue":"Archives of Pathology & Laboratory Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Medicine; Terminology; Identifier; Selection (genetic algorithm); Computer science; Medical physics; Artificial intelligence; Programming language","score_opus":0.03777631581919645,"score_gpt":0.31480496834098876,"score_spread":0.2770286525217923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2980234588","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9989888,0.000062507956,0.00019618319,0.00018709489,0.0000034214363,0.000047979876,0.00015637313,0.000005555976,0.00035209078],"genre_scores_gemma":[0.9988625,0.00012268289,0.0004141733,0.00018786664,0.000005625412,0.0000517106,0.00018353292,0.000003941218,0.00016800474],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9901827,0.003461836,0.0011297304,0.0011155735,0.00318389,0.0009263342],"domain_scores_gemma":[0.96235377,0.0116483215,0.013733543,0.0014620298,0.008133302,0.0026690282],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0120428195,0.00021839756,0.00023847484,0.0016815629,0.0009148569,0.0009347363,0.0010091527,0.000725483,0.0012769743],"category_scores_gemma":[0.03489413,0.00037110213,0.00026720294,0.0016683978,0.0013259965,0.00096931774,0.0011790757,0.00070828595,0.00028367367],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055983463,0.00007335625,0.9875853,0.000025834168,0.000008604785,0.00006297908,0.0061512724,0.000027990836,0.00048249838,0.0000400641,0.00046900357,0.005017062],"study_design_scores_gemma":[0.000007769529,0.00032155018,0.9842821,0.00004605438,0.000006729034,0.00027080753,0.012915624,0.000417448,0.00030548227,0.000020429998,0.001390964,0.000015055203],"about_ca_topic_score_codex":0.060312368,"about_ca_topic_score_gemma":0.047012705,"teacher_disagreement_score":0.9879572,"about_ca_system_score_codex":0.0022226083,"about_ca_system_score_gemma":0.0037913297,"threshold_uncertainty_score":0.11992258},"labels":[],"label_agreement":null},{"id":"W2980729377","doi":"10.1002/pra2.199","title":"Developing community phenotype ontologies: Understanding users' preferences","year":2019,"lang":"en","type":"article","venue":"Proceedings of the Association for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Agriculture and Agri-Food Canada; University of Manitoba","funders":"National Science Foundation","keywords":"Wizard; Ontology; Computer science; World Wide Web; Phenotype; Information retrieval; Biology; Epistemology; Genetics","score_opus":0.03753585336760046,"score_gpt":0.27786760165272567,"score_spread":0.2403317482851252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2980729377","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.889875,0.00020694964,0.09602564,0.00086069066,0.000026926416,0.00034537114,0.0015611856,0.0043561994,0.006742052],"genre_scores_gemma":[0.9259901,0.000106775704,0.069406316,0.0001625473,0.000010732153,0.0001994987,0.0021734051,0.00042173904,0.0015288716],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.99562293,0.0024304034,0.00037181622,0.0005381321,0.00085985137,0.0001767641],"domain_scores_gemma":[0.96010375,0.031794332,0.0010410834,0.0024497763,0.0039549354,0.0006561405],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009610115,0.00058364496,0.00051391753,0.0015669216,0.00074162567,0.0024348511,0.0008195083,0.00086062134,0.0022867238],"category_scores_gemma":[0.04283026,0.00034776362,0.0006176646,0.00088441843,0.0003404141,0.0055992627,0.0019357475,0.0006633717,0.0006856158],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038669256,0.0017530823,0.32606646,0.0016485378,0.00036520627,0.0012116142,0.057079382,0.006967629,0.0578671,0.007913038,0.016892461,0.5183686],"study_design_scores_gemma":[0.00065769046,0.0019198892,0.24103822,0.00067799335,0.00080676645,0.0022570905,0.060096525,0.48175982,0.06664511,0.027839253,0.11565256,0.00064913137],"about_ca_topic_score_codex":0.004073367,"about_ca_topic_score_gemma":0.005994377,"teacher_disagreement_score":0.009610115,"about_ca_system_score_codex":0.000531576,"about_ca_system_score_gemma":0.00051249453,"threshold_uncertainty_score":0.050823748},"labels":[],"label_agreement":null},{"id":"W2981124990","doi":"","title":"Développement d'outils bio-informatiques pour l'analyse de données épigénomiques avec référence externe et pour l’évaluation du nombre de couples à risque à partir de fichiers de variants génétiques","year":2018,"lang":"fr","type":"article","venue":"Knowledge UdeS (Institutional Deposit of the University of Sherbrooke)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of General Medical Sciences; McGill University; Génome Québec; Fonds de Recherche du Québec - Santé; Compute Canada; Natural Sciences and Engineering Research Council of Canada","keywords":"Humanities; Philosophy","score_opus":0.025404100226225877,"score_gpt":0.25544893532082313,"score_spread":0.23004483509459725,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2981124990","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06906,0.0021241526,0.88888466,0.0008458438,0.0004326588,0.0006505336,0.010113255,0.024402622,0.00348628],"genre_scores_gemma":[0.1046133,0.0014454736,0.8634295,0.00063100265,0.00009406499,0.0014751939,0.017257405,0.0026308685,0.008423167],"study_design_codex":"bench_or_experimental","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917277,0.0022889292,0.00073507405,0.0018189563,0.003044224,0.00038510628],"domain_scores_gemma":[0.97092193,0.018325865,0.0015429591,0.0038638688,0.004913527,0.00043192448],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009416703,0.0020681147,0.0017953046,0.0028409287,0.00077248557,0.0037503915,0.0016879259,0.00219191,0.007144984],"category_scores_gemma":[0.022748373,0.0012568969,0.0022863562,0.0019063907,0.0010433741,0.0022973674,0.0015262957,0.0022932452,0.004158633],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017978754,0.0004994675,0.015131756,0.0026649395,0.0006425196,0.0007064286,0.001552612,0.034107517,0.50951666,0.0052843583,0.00930114,0.41879472],"study_design_scores_gemma":[0.00016245758,0.0007707732,0.015601779,0.00030775453,0.00033966667,0.00093932584,0.000542285,0.13448182,0.7599922,0.005308545,0.0812071,0.00034639376],"about_ca_topic_score_codex":0.004947143,"about_ca_topic_score_gemma":0.006823668,"teacher_disagreement_score":0.009416703,"about_ca_system_score_codex":0.0015789061,"about_ca_system_score_gemma":0.002721771,"threshold_uncertainty_score":0.049800873},"labels":[],"label_agreement":null},{"id":"W2982546097","doi":"10.3390/electronics8111235","title":"A Review of Automatic Phenotyping Approaches using Electronic Health Records","year":2019,"lang":"en","type":"review","venue":"Electronics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":58,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Charles Darwin University; Trent University; Nottingham Trent University","keywords":"Computer science; Popularity; Biomedical text mining; Information retrieval; Artificial intelligence; Machine learning; Health records; Subject (documents); Data science; Information extraction; Electronic health record; Natural language processing; Health care; World Wide Web; Text mining","score_opus":0.11076840308400433,"score_gpt":0.37023924153910004,"score_spread":0.2594708384550957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2982546097","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00022649068,0.996822,0.0011549707,0.0003813476,0.00013298355,0.000027314454,0.00012028599,0.00002753635,0.0011070238],"genre_scores_gemma":[0.0010522116,0.9962089,0.0019834256,0.00024901715,0.00008488783,0.000028475924,0.00015734893,0.000005759555,0.00023003826],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.99848795,0.00038698796,0.00036032838,0.00025483713,0.00046103672,0.00004887388],"domain_scores_gemma":[0.9931559,0.004933572,0.0005273853,0.00015244639,0.0011489883,0.00008157269],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030734506,0.0010466983,0.0013729157,0.008364784,0.00040437412,0.0013039967,0.0015104401,0.0012193628,0.0036726627],"category_scores_gemma":[0.007864911,0.00048492002,0.0014930853,0.007693798,0.00058201817,0.0022765659,0.0008385062,0.0009687883,0.0017658427],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004917913,0.000043095548,0.00046769722,0.066750534,0.00020963783,0.00012869318,0.00014346346,0.00031671018,0.00071295514,0.0018613951,0.012730368,0.9165862],"study_design_scores_gemma":[0.000024947523,0.00015388857,0.005299888,0.07117847,0.0012987396,0.0019400925,0.00031680157,0.00045353902,0.0016735788,0.0031028753,0.9144715,0.000085748],"about_ca_topic_score_codex":0.0026928922,"about_ca_topic_score_gemma":0.0032838273,"teacher_disagreement_score":0.008364784,"about_ca_system_score_codex":0.0009072786,"about_ca_system_score_gemma":0.0029878004,"threshold_uncertainty_score":0.016254187},"labels":[],"label_agreement":null},{"id":"W2983608826","doi":"10.18653/v1/d19-6219","title":"Recognizing UMLS Semantic Types with Deep Learning","year":2019,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Privacy Analytics (Canada); National Research Council Canada","funders":"","keywords":"Unified Medical Language System; Computer science; Deep learning; Natural language processing; Information retrieval; Artificial intelligence","score_opus":0.007885366663328479,"score_gpt":0.23021598848755392,"score_spread":0.22233062182422544,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2983608826","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0502329,0.0033223992,0.869588,0.0029403474,0.0010073317,0.00040641968,0.017578907,0.042790044,0.012133458],"genre_scores_gemma":[0.24388248,0.0014704143,0.7001455,0.0007354217,0.00017808456,0.00029683774,0.04404553,0.0011539367,0.008091789],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985209,0.00025373534,0.00020578236,0.00038692748,0.00046051145,0.00017211773],"domain_scores_gemma":[0.9967493,0.0014433115,0.0003000561,0.00054238894,0.00081106456,0.00015393582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012875876,0.0014851006,0.000879807,0.0069371862,0.000882754,0.0025952554,0.0016890109,0.001480726,0.004998498],"category_scores_gemma":[0.0060100756,0.0009078346,0.00256453,0.0035805185,0.0006532064,0.0072865244,0.0030478856,0.0019157787,0.0038545048],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040434496,0.00037436187,0.0082605835,0.0006192404,0.00014712654,0.00058011094,0.0006328305,0.015807616,0.0069583612,0.04022052,0.08719953,0.83879536],"study_design_scores_gemma":[0.0000785069,0.00007561413,0.0022511415,0.00043696855,0.00012440272,0.00062791054,0.0010031513,0.67827547,0.012889233,0.20464341,0.09951149,0.00008271264],"about_ca_topic_score_codex":0.0136660375,"about_ca_topic_score_gemma":0.034157712,"teacher_disagreement_score":0.0136660375,"about_ca_system_score_codex":0.0016974146,"about_ca_system_score_gemma":0.0017697404,"threshold_uncertainty_score":0.027172983},"labels":[],"label_agreement":null},{"id":"W2989049020","doi":"10.1145/3357384.3358128","title":"Health Card Retrieval for Consumer Health Search","year":2019,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute of Automation, Chinese Academy of Sciences; University of California, San Diego; Wuhan University; Georgetown University; Technische Universität Berlin; Shanghai Jiao Tong University; Zhejiang University; Electronics and Telecommunications Research Institute; RMIT University; Tsinghua University; York University; Harbin University of Science and Technology; Harbin Institute of Technology; University of Illinois at Urbana-Champaign; Northeast Forestry University; Chinese Academy of Sciences; Tencent; Technische Universiteit Delft; Case Western Reserve University; Nanjing University; Microsoft Research; Lembaga Pengelola Dana Pendidikan; Università degli Studi di Udine; National University of Defense Technology; Arizona State University; University of North Carolina at Chapel Hill; Nanjing University of Aeronautics and Astronautics; Pennsylvania State University; Microsoft Research Asia; Peking University","keywords":"Computer science; Information retrieval; World Wide Web; Internet privacy","score_opus":0.03126670483852918,"score_gpt":0.35211440262303445,"score_spread":0.3208476977845053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2989049020","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16472717,0.013819452,0.73618966,0.004783556,0.0005511975,0.003162769,0.014233539,0.011713898,0.050818834],"genre_scores_gemma":[0.3903539,0.001996911,0.5846472,0.0006020253,0.0003712037,0.00041114746,0.0112380665,0.000277409,0.010102177],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9960432,0.0016805838,0.00043812772,0.0004895879,0.001163997,0.0001845629],"domain_scores_gemma":[0.9910767,0.005612716,0.0006019108,0.0011144418,0.0013825925,0.00021166953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004113319,0.0008635087,0.0010264458,0.009065474,0.0008312335,0.0026450055,0.0012237511,0.0014734245,0.008711615],"category_scores_gemma":[0.020158267,0.00022523888,0.00085676846,0.0051667676,0.0005655616,0.0051718634,0.0011619797,0.00097300275,0.0043940926],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046080904,0.0005586736,0.013575956,0.0012737829,0.00018519783,0.000157945,0.00055900513,0.01185784,0.0064446875,0.019019514,0.040810667,0.9050959],"study_design_scores_gemma":[0.00029471575,0.0010653151,0.032552246,0.00043252046,0.00035225038,0.0020437934,0.0019626142,0.7533515,0.022321815,0.057665795,0.12766553,0.0002918407],"about_ca_topic_score_codex":0.005957809,"about_ca_topic_score_gemma":0.009623799,"teacher_disagreement_score":0.009065474,"about_ca_system_score_codex":0.0014089168,"about_ca_system_score_gemma":0.0014776798,"threshold_uncertainty_score":0.029143214},"labels":[],"label_agreement":null},{"id":"W2989997401","doi":"10.3233/efi-190338","title":"Towards collective intelligence in a national community of physicians","year":2019,"lang":"en","type":"article","venue":"Education for Information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Collective intelligence; Constructive; Public relations; Psychology; Knowledge management; Medical education; Work (physics); Information sharing; Sociology; Computer science; Medicine; Political science; World Wide Web; Engineering","score_opus":0.019193826723586636,"score_gpt":0.31814045721128037,"score_spread":0.2989466304876937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2989997401","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14949304,0.0016900321,0.4555083,0.16352724,0.0014692624,0.0010201285,0.00025654797,0.0014762386,0.22555925],"genre_scores_gemma":[0.6687631,0.00080624415,0.27764657,0.009634156,0.0008094274,0.0006207617,0.00040228697,0.00020963768,0.041107796],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98126376,0.0106221745,0.0007220539,0.0022537971,0.004004021,0.0011342085],"domain_scores_gemma":[0.9443705,0.023710243,0.004613678,0.008095234,0.010374551,0.008835738],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03321807,0.00044557275,0.00057447015,0.004679168,0.011545023,0.0135486135,0.0021385786,0.0037940594,0.006897734],"category_scores_gemma":[0.060530435,0.000706092,0.0010128248,0.0030681258,0.009815469,0.01433486,0.023386832,0.0036988633,0.0014882407],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016997744,0.00036084227,0.024765247,0.0003254879,0.00020309653,0.0009243457,0.08193328,0.0047631543,0.0032265787,0.6177301,0.043745898,0.22185203],"study_design_scores_gemma":[0.0001123739,0.000100183504,0.006141361,0.00037206715,0.00008931366,0.0003126362,0.026135404,0.024074584,0.0012866928,0.5663539,0.37489402,0.00012742232],"about_ca_topic_score_codex":0.02017773,"about_ca_topic_score_gemma":0.019221142,"teacher_disagreement_score":0.03321807,"about_ca_system_score_codex":0.0054782145,"about_ca_system_score_gemma":0.01871591,"threshold_uncertainty_score":0.17567605},"labels":[],"label_agreement":null},{"id":"W2991568540","doi":"10.1016/j.procs.2019.11.079","title":"COMPETENCY QUESTIONS FOR BIOMEDICAL ONTOLOGY REUSE","year":2019,"lang":"en","type":"article","venue":"Procedia Computer Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Computer science; Reuse; Ontology; Interoperability; Domain (mathematical analysis); Scope (computer science); Process (computing); Ontology engineering; Upper ontology; Process ontology; Open Biomedical Ontologies; Software engineering; Data science; Semantics (computer science); Semantic interoperability; Knowledge management; Domain knowledge; World Wide Web; Suggested Upper Merged Ontology; Programming language","score_opus":0.010720130646637476,"score_gpt":0.2784180533142503,"score_spread":0.26769792266761283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2991568540","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037543487,0.000924271,0.9241932,0.016260812,0.0001703449,0.0005322301,0.00025874216,0.00028366118,0.019833224],"genre_scores_gemma":[0.4673743,0.00057303195,0.52444875,0.0020763373,0.0002409663,0.0006883604,0.0007047147,0.00015775867,0.0037357914],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9592085,0.025154782,0.0042108498,0.0037889034,0.0063248854,0.00131219],"domain_scores_gemma":[0.8889947,0.07468751,0.0061964244,0.012133802,0.015674016,0.0023135582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034131616,0.0010166096,0.00089053944,0.0055547194,0.003347492,0.0049391384,0.0022383374,0.004363256,0.0029287348],"category_scores_gemma":[0.13604872,0.00074044464,0.0024063922,0.0024919666,0.0122719975,0.018833993,0.010508981,0.0046036337,0.0006402102],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000052276293,0.00013516059,0.0036547664,0.00032340857,0.00006875602,0.0005666413,0.0074601998,0.005956188,0.0013480132,0.89852136,0.0030040143,0.07890918],"study_design_scores_gemma":[0.000028876702,0.000041915304,0.0012311349,0.00034932565,0.00004537587,0.00050556177,0.0032642488,0.02580732,0.0021584474,0.9342074,0.032281365,0.000079019424],"about_ca_topic_score_codex":0.0097543765,"about_ca_topic_score_gemma":0.00400629,"teacher_disagreement_score":0.034131616,"about_ca_system_score_codex":0.004244849,"about_ca_system_score_gemma":0.0061675655,"threshold_uncertainty_score":0.18050742},"labels":[],"label_agreement":null},{"id":"W2993172114","doi":"10.1186/s13073-019-0686-y","title":"Text-mining clinically relevant cancer biomarkers for curation into the CIViC database","year":2019,"lang":"en","type":"review","venue":"Genome Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":46,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"National Human Genome Research Institute; National Institutes of Health; National Cancer Institute; Compute Canada","keywords":"Data curation; Medicine; Computational biology; Data science; Bioinformatics; Computer science; Biology","score_opus":0.1005962690517609,"score_gpt":0.4202753033606965,"score_spread":0.31967903430893563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2993172114","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021966398,0.009718816,0.05372967,0.004067258,0.00050622644,0.0031252778,0.8778704,0.018581616,0.010434366],"genre_scores_gemma":[0.028775908,0.0034408453,0.19676441,0.0013552863,0.00020813651,0.0024382467,0.7648465,0.00075237517,0.0014182846],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9946169,0.0010143286,0.0018449038,0.0012605989,0.0010904653,0.00017274826],"domain_scores_gemma":[0.9618858,0.021703804,0.0056297746,0.0033127614,0.006384221,0.0010836395],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008839363,0.0020717306,0.0022331479,0.03034797,0.0016511164,0.0039031953,0.0030829487,0.0028704526,0.011132637],"category_scores_gemma":[0.044106018,0.0009383846,0.0020131078,0.016847236,0.0007335603,0.0037597194,0.0036754396,0.0020856853,0.006380804],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013687629,0.0006730762,0.0259172,0.059231773,0.001354217,0.006230119,0.0032877484,0.0072753127,0.034962576,0.013865068,0.53742707,0.30840707],"study_design_scores_gemma":[0.0005954444,0.00046048113,0.034129236,0.011063336,0.002204722,0.0029070654,0.002252593,0.030025283,0.029872078,0.015975118,0.8701529,0.00036171233],"about_ca_topic_score_codex":0.005463955,"about_ca_topic_score_gemma":0.01091381,"teacher_disagreement_score":0.03034797,"about_ca_system_score_codex":0.0025771298,"about_ca_system_score_gemma":0.00862782,"threshold_uncertainty_score":0.046747625},"labels":[],"label_agreement":null},{"id":"W2994106476","doi":"10.22148/16.058","title":"Annotating Narrative Levels: Review of Guideline No. 5","year":2019,"lang":"en","type":"article","venue":"Journal of Cultural Analytics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Narrative; Guideline; Narrative criticism; Narrative history; Narrative network; Narrative inquiry; Linguistics; Psychology; Computer science; Literature; History; Philosophy; Art; Political science","score_opus":0.030970730186825144,"score_gpt":0.34129124731195526,"score_spread":0.3103205171251301,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2994106476","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004531219,0.6414632,0.11695241,0.1459507,0.030542651,0.0040322067,0.009816903,0.0019018571,0.044808943],"genre_scores_gemma":[0.028217653,0.5252422,0.26903903,0.10772503,0.0053278184,0.009453668,0.028442249,0.00140085,0.025151398],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97928303,0.0068663373,0.0068434365,0.0015570718,0.0048280223,0.0006221345],"domain_scores_gemma":[0.9001173,0.04796171,0.0058935834,0.005381677,0.03901658,0.0016292573],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035922896,0.00095822325,0.0021457889,0.012126891,0.0020846922,0.004505759,0.00765979,0.00590927,0.005353027],"category_scores_gemma":[0.1093183,0.0012879947,0.0031991652,0.0069087995,0.003446407,0.006685216,0.0052415933,0.005058719,0.005287352],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013684633,0.0000990625,0.0011062915,0.07596684,0.00026667025,0.0004931872,0.0036608607,0.00063425,0.002523995,0.029729843,0.3936764,0.4917058],"study_design_scores_gemma":[0.000022132794,0.000035155572,0.0010363802,0.07177406,0.00026770533,0.0003984152,0.00054774154,0.00017373083,0.0008474672,0.0060363263,0.91881496,0.000045930465],"about_ca_topic_score_codex":0.02287276,"about_ca_topic_score_gemma":0.046304975,"teacher_disagreement_score":0.035922896,"about_ca_system_score_codex":0.005498651,"about_ca_system_score_gemma":0.024455769,"threshold_uncertainty_score":0.18998075},"labels":[],"label_agreement":null},{"id":"W2995438613","doi":"10.48550/arxiv.1912.06174","title":"Training without training data: Improving the generalizability of automated medical abbreviation disambiguation","year":2019,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Generalizability theory; Computer science; Context (archaeology); Artificial intelligence; Training set; Natural language processing; Labeled data; Machine learning; Representation (politics); Training (meteorology); Scarcity; Psychology","score_opus":0.18410751469812822,"score_gpt":0.2660201738176082,"score_spread":0.08191265911947995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2995438613","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31246436,0.0057255416,0.6390265,0.005226135,0.0006442094,0.000591848,0.0054423693,0.024447214,0.0064317933],"genre_scores_gemma":[0.7852929,0.0010602808,0.19468506,0.0018182643,0.00037855154,0.0003223762,0.013260648,0.00066014077,0.0025217582],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99448,0.0023357978,0.0004060127,0.0019550158,0.0006025837,0.00022051993],"domain_scores_gemma":[0.9793251,0.013241688,0.0008594016,0.0047069704,0.0015800274,0.00028688807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010029906,0.0021049299,0.001686807,0.0023923453,0.00093787967,0.002097509,0.0029580123,0.0025671585,0.0015740426],"category_scores_gemma":[0.034316283,0.00078417594,0.0016285798,0.0023911155,0.0012091041,0.004556637,0.0033360522,0.0034921118,0.0015322741],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013916756,0.0006313268,0.039104268,0.00051134254,0.00060957164,0.0006708653,0.0006528673,0.27230448,0.010894734,0.0026462027,0.028395541,0.6421872],"study_design_scores_gemma":[0.00008570204,0.0001309205,0.004585824,0.00008007964,0.000119142875,0.00017198414,0.0001463777,0.9758301,0.005225713,0.0069819842,0.0065998244,0.00004226643],"about_ca_topic_score_codex":0.013103871,"about_ca_topic_score_gemma":0.014180676,"teacher_disagreement_score":0.013103871,"about_ca_system_score_codex":0.0010893024,"about_ca_system_score_gemma":0.002064189,"threshold_uncertainty_score":0.053043902},"labels":[],"label_agreement":null},{"id":"W2996406085","doi":"10.2196/16042","title":"Clinical Annotation Research Kit (CLARK): Computable Phenotyping Using Machine Learning","year":2019,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences","keywords":"Machine learning; Artificial intelligence; Computer science; Naive Bayes classifier; Random forest; Support vector machine; Annotation; Decision tree; Classifier (UML); Natural language processing","score_opus":0.08154325136589928,"score_gpt":0.4319228986456122,"score_spread":0.3503796472797129,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996406085","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061868513,0.0009192278,0.39495975,0.004539446,0.0007547917,0.0019003839,0.11686219,0.4445314,0.029345939],"genre_scores_gemma":[0.046430025,0.001266348,0.6821088,0.0055445316,0.00035684323,0.005463832,0.16482455,0.05844589,0.035559185],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953826,0.0014273181,0.0008201441,0.0009942917,0.0011954225,0.00018024715],"domain_scores_gemma":[0.96515733,0.024328185,0.0019917923,0.0033814143,0.0043594204,0.0007817512],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007237319,0.0021850348,0.0009919835,0.00410626,0.0005895085,0.0031383054,0.0032041748,0.0019804726,0.08670174],"category_scores_gemma":[0.04900406,0.0016283388,0.0012216717,0.002844274,0.0008329799,0.0041669686,0.005125122,0.002060146,0.04545108],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006953849,0.00010305197,0.0035202806,0.0025892379,0.00010753024,0.0008787623,0.0009573688,0.0015531095,0.004239918,0.015013665,0.7655603,0.20478141],"study_design_scores_gemma":[0.0007279264,0.00024124492,0.0058480105,0.0013479765,0.00011713512,0.0021873093,0.0003929414,0.023884116,0.018362485,0.047906328,0.8986735,0.00031101794],"about_ca_topic_score_codex":0.0024946956,"about_ca_topic_score_gemma":0.003402713,"teacher_disagreement_score":0.08670174,"about_ca_system_score_codex":0.001136419,"about_ca_system_score_gemma":0.0029750452,"threshold_uncertainty_score":0.29004627},"labels":[],"label_agreement":null},{"id":"W2996617694","doi":"10.2196/13498","title":"Adverse Childhood Experiences Ontology for Mental Health Surveillance, Research, and Evaluation: Advanced Knowledge Representation and Semantic Web Techniques","year":2019,"lang":"en","type":"article","venue":"JMIR Mental Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Ontology; Mental health; Semantic Web; Semantic network; Knowledge representation and reasoning; Representation (politics); Mental representation; Semantics (computer science)","score_opus":0.03929296848855788,"score_gpt":0.4402923995142207,"score_spread":0.4009994310256628,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996617694","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014711719,0.0011897982,0.94860137,0.0036737758,0.00014554938,0.002787933,0.010174763,0.0057644043,0.012950653],"genre_scores_gemma":[0.097187385,0.0014580457,0.8801834,0.00061304367,0.000039594157,0.0018305405,0.01706711,0.000272147,0.0013488074],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98939663,0.003972675,0.0020711909,0.0009333251,0.0032848176,0.00034141354],"domain_scores_gemma":[0.9865682,0.0069359574,0.0012667851,0.0024526,0.002285451,0.00049097586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013265078,0.00077645644,0.000775999,0.00643522,0.0012074036,0.005730095,0.0022743498,0.0015345059,0.0021549084],"category_scores_gemma":[0.027555194,0.00050633756,0.00253824,0.0056107203,0.001489938,0.0073754922,0.005806343,0.0020257332,0.00059342274],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005703712,0.001330565,0.0119046,0.0046958015,0.00081519585,0.001406139,0.0042589107,0.04353218,0.009243968,0.39355862,0.035154074,0.49352956],"study_design_scores_gemma":[0.00024597603,0.00027227038,0.008945532,0.0034660653,0.0007438786,0.0010929788,0.0036111115,0.3102536,0.014558524,0.3330023,0.3235748,0.00023299984],"about_ca_topic_score_codex":0.021277186,"about_ca_topic_score_gemma":0.021246435,"teacher_disagreement_score":0.021277186,"about_ca_system_score_codex":0.004113778,"about_ca_system_score_gemma":0.007473136,"threshold_uncertainty_score":0.07015324},"labels":[],"label_agreement":null},{"id":"W2996653911","doi":"10.1186/s12920-019-0637-x","title":"GTX.Digest.VCF: an online NGS data interpretation system based on intelligent gene ranking and large-scale text mining","year":2019,"lang":"en","type":"article","venue":"BMC Medical Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Royal University Hospital","funders":"Key Technologies Research and Development Program; National Natural Science Foundation of China","keywords":"Computer science; Workflow; Annotation; Genomics; Data mining; Biomedical text mining; Data science; Computational biology; Bioinformatics; Information retrieval; Gene; Text mining; Genome; Artificial intelligence; Biology; Genetics; Database","score_opus":0.03590114342688194,"score_gpt":0.2910567341227951,"score_spread":0.25515559069591315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2996653911","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019208152,0.0009072783,0.20042624,0.001110756,0.00035882747,0.001119615,0.18991496,0.58159566,0.0053584664],"genre_scores_gemma":[0.07491735,0.00067061867,0.5776772,0.0010103218,0.00025446643,0.0018699423,0.3130864,0.019787695,0.010726012],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99883634,0.00018716321,0.0001872682,0.000375955,0.0003397259,0.000073511736],"domain_scores_gemma":[0.99555236,0.0022439747,0.00059567246,0.00064790924,0.0007135055,0.00024653831],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035457218,0.0024033322,0.0011586248,0.0073520183,0.0007595272,0.001965345,0.0016682699,0.0013962707,0.024637293],"category_scores_gemma":[0.009596406,0.00079095055,0.0015428065,0.0033542104,0.0004223023,0.0020053345,0.001754739,0.0008595147,0.011465841],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018140869,0.00027801728,0.015785294,0.0019794775,0.0005606768,0.001671376,0.00071864564,0.006332999,0.027807433,0.0026559518,0.48840934,0.4519867],"study_design_scores_gemma":[0.0022731787,0.00079127616,0.052422497,0.00097201637,0.00046842408,0.0057108095,0.0007055605,0.33776954,0.08802972,0.028071366,0.48197603,0.0008095256],"about_ca_topic_score_codex":0.005743337,"about_ca_topic_score_gemma":0.005030685,"teacher_disagreement_score":0.024637293,"about_ca_system_score_codex":0.001073178,"about_ca_system_score_gemma":0.0023516717,"threshold_uncertainty_score":0.08241993},"labels":[],"label_agreement":null},{"id":"W2997710399","doi":"10.1109/vahc47919.2019.8945033","title":"Comparing ICD-Data Across Countries: A Case for Visualization?","year":2019,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Comparability; Generalizability theory; Visual analytics; Visualization; Coding (social sciences); Data science; Data visualization; Computer science; Analytics; Data collection; Data mining; Statistics","score_opus":0.0631661071261037,"score_gpt":0.39190811514664575,"score_spread":0.32874200802054204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2997710399","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16380905,0.032273732,0.33324155,0.3693551,0.007282499,0.0011630023,0.016531046,0.002436959,0.07390711],"genre_scores_gemma":[0.7685064,0.0036662281,0.20416136,0.01343698,0.0013981425,0.0012891057,0.004594522,0.0014282796,0.0015190007],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.740061,0.21488455,0.015824437,0.009802498,0.017045343,0.0023822784],"domain_scores_gemma":[0.43202952,0.42409396,0.036153734,0.060401794,0.04416398,0.0031569493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.18665616,0.0011672771,0.0017354682,0.013233042,0.0026656685,0.017777126,0.0038760093,0.0027958958,0.0062289946],"category_scores_gemma":[0.53670114,0.00088769285,0.0017011418,0.026329253,0.0074960045,0.024418,0.013088986,0.0057841665,0.0008369221],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009279586,0.00015529811,0.12984025,0.00581383,0.0018199469,0.0012713106,0.09613823,0.0037319388,0.001000161,0.2993161,0.07982761,0.38015732],"study_design_scores_gemma":[0.0002038075,0.0002622389,0.086760364,0.0147773605,0.0005432954,0.002630428,0.09115087,0.010928109,0.002108331,0.48240623,0.30778447,0.0004445543],"about_ca_topic_score_codex":0.010074224,"about_ca_topic_score_gemma":0.008349227,"teacher_disagreement_score":0.18665616,"about_ca_system_score_codex":0.0047964617,"about_ca_system_score_gemma":0.0044205077,"threshold_uncertainty_score":0.98714393},"labels":[],"label_agreement":null},{"id":"W2998584897","doi":"10.35502/jcswb.114","title":"Bringing research closer to collaborative practice at LEPH2019","year":2019,"lang":"en","type":"article","venue":"Journal of Community Safety and Well-Being","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Sociology","score_opus":0.018975492248721273,"score_gpt":0.3460688825136306,"score_spread":0.32709339026490936,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2998584897","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.072487354,0.0032510331,0.07109405,0.4915789,0.026518222,0.0022736879,0.015263108,0.014690421,0.3028432],"genre_scores_gemma":[0.24658372,0.000969264,0.1280059,0.029706327,0.006623212,0.0014969666,0.015975457,0.0060339263,0.5646052],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9802817,0.009226593,0.00047864718,0.0016761766,0.0056351465,0.002701668],"domain_scores_gemma":[0.9302979,0.011611135,0.0015282292,0.0063864267,0.0091753015,0.04100104],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03324294,0.00047462076,0.000678754,0.0028078156,0.006917432,0.016834991,0.0027281747,0.0059227394,0.095139295],"category_scores_gemma":[0.040216774,0.0005855472,0.0009600889,0.0026292559,0.0033555226,0.009423186,0.017091958,0.0071729687,0.025767049],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005458709,0.00079897756,0.0071631535,0.0003895886,0.00006079441,0.0010852527,0.011979518,0.0007887443,0.0035179954,0.0387353,0.70141375,0.23352112],"study_design_scores_gemma":[0.00008066672,0.00013867398,0.0074736485,0.00023504153,0.000007697314,0.0001406359,0.0050867232,0.0015128754,0.0012857878,0.02076155,0.9632197,0.000057145986],"about_ca_topic_score_codex":0.011135913,"about_ca_topic_score_gemma":0.03161458,"teacher_disagreement_score":0.095139295,"about_ca_system_score_codex":0.009765858,"about_ca_system_score_gemma":0.021706661,"threshold_uncertainty_score":0.31827265},"labels":[],"label_agreement":null},{"id":"W3007008050","doi":"","title":"Guides: Vancouver reference style (based on Citing Medicine): Multimedia formats","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Multimedia; Style (visual arts); Computer science; World Wide Web; Visual arts; Art","score_opus":0.048426762152045756,"score_gpt":0.29627229004891265,"score_spread":0.2478455278968669,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3007008050","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0012469187,0.00061562017,0.08519428,0.0027151962,0.0010106241,0.0007266951,0.44580907,0.12040811,0.34227344],"genre_scores_gemma":[0.0063778157,0.0015961305,0.118783794,0.0008320394,0.00043608306,0.00086201617,0.429516,0.078917466,0.36267865],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9982628,0.0002371646,0.00036206347,0.00021646803,0.0008220476,0.00009945273],"domain_scores_gemma":[0.9840655,0.004036251,0.00053986214,0.0018504001,0.008712576,0.00079537294],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002267814,0.0020018874,0.0013707535,0.012776606,0.0016911816,0.0082313595,0.0035114435,0.0024771423,0.52428967],"category_scores_gemma":[0.022173008,0.0016909973,0.0008882178,0.02193891,0.00072974496,0.006210135,0.0029843184,0.0020340057,0.4278588],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036455913,0.000020163365,0.00009441942,0.00034478414,0.0000064372516,0.000033996796,0.00012046589,0.00017987409,0.00028706784,0.0041358215,0.95239514,0.042345293],"study_design_scores_gemma":[0.00001838607,0.000006142889,0.00028855246,0.00015073163,0.0000073752926,0.00007340352,0.00006519983,0.0005413801,0.00075069745,0.0026359723,0.99543875,0.000023448736],"about_ca_topic_score_codex":0.04620641,"about_ca_topic_score_gemma":0.08552556,"teacher_disagreement_score":0.99176866,"about_ca_system_score_codex":0.002301507,"about_ca_system_score_gemma":0.004828853,"threshold_uncertainty_score":0.67854303},"labels":[],"label_agreement":null},{"id":"W3007410695","doi":"10.1136/jclinpath-2019-206370","title":"Putting the patient at the centre of pathology: an innovative approach to patient education—MyPathologyReport.ca","year":2020,"lang":"en","type":"review","venue":"Journal of Clinical Pathology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; University of Ottawa","funders":"","keywords":"Medical diagnosis; Medicine; Variety (cybernetics); Patient care; Patient education; MEDLINE; Medical education; Pathology; Computer science; Nursing","score_opus":0.09021415037588325,"score_gpt":0.4252743902464943,"score_spread":0.33506023987061107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3007410695","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024983084,0.010457722,0.054225165,0.4939684,0.019451661,0.0018860972,0.0022937148,0.024127634,0.36860648],"genre_scores_gemma":[0.14000435,0.024724085,0.28847852,0.27661192,0.026414433,0.0032122093,0.0038490843,0.00534587,0.23135945],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99675167,0.0016178911,0.00026768536,0.00026907952,0.0008429539,0.0002507234],"domain_scores_gemma":[0.9746591,0.008134037,0.0021460275,0.0022473964,0.0024385143,0.010374962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005869925,0.00060222985,0.000438843,0.0019198714,0.0020783886,0.006750573,0.0025849899,0.003730848,0.059486486],"category_scores_gemma":[0.021448256,0.00053668447,0.000709115,0.0012447383,0.0017402563,0.009878689,0.009804154,0.0072934334,0.030955698],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000032164007,0.0003153951,0.002099425,0.00028473634,0.00001117631,0.0005805179,0.0011699523,0.000038058188,0.0003869911,0.001978409,0.75479877,0.23830444],"study_design_scores_gemma":[0.00004770007,0.00017060431,0.0033041928,0.0006525957,0.000015427322,0.0033035628,0.0017451689,0.0002261494,0.00039260273,0.0035890918,0.98649466,0.000058296166],"about_ca_topic_score_codex":0.0005340447,"about_ca_topic_score_gemma":0.0021266514,"teacher_disagreement_score":0.99946594,"about_ca_system_score_codex":0.0009946858,"about_ca_system_score_gemma":0.0036027026,"threshold_uncertainty_score":0.19900209},"labels":[],"label_agreement":null},{"id":"W3008294444","doi":"10.2196/16777","title":"Translating Clinical Questions by Physicians Into Searchable Queries: Analytical Survey Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"McMaster University; Impact","funders":"","keywords":"Computer science; Information retrieval; Context (archaeology); Service (business); MEDLINE; Search engine; World Wide Web; Medicine","score_opus":0.046944464602369826,"score_gpt":0.43421166438750075,"score_spread":0.3872671997851309,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3008294444","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9896327,0.0016137192,0.0014763746,0.0006875084,0.00001738227,0.0018044072,0.0032184387,0.000054892458,0.0014945279],"genre_scores_gemma":[0.98756325,0.0014033525,0.003958449,0.0014642813,0.000041843814,0.0031604413,0.0020433841,0.000041853713,0.00032329446],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.94830734,0.025443278,0.01391926,0.004454913,0.0058142543,0.0020609987],"domain_scores_gemma":[0.7219966,0.20055169,0.04739302,0.006053721,0.020266995,0.0037378643],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.034819234,0.00043076795,0.0012939274,0.006270909,0.0011176462,0.0027660185,0.0010097971,0.0013602847,0.0030322701],"category_scores_gemma":[0.21038975,0.00077676214,0.0011269393,0.008852278,0.0016783025,0.0041013723,0.002576728,0.0012367624,0.0017328416],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015452463,0.0014088118,0.8952802,0.004220152,0.00048385593,0.0004320339,0.042900227,0.0003454406,0.00063044694,0.00057471596,0.0058121416,0.046366766],"study_design_scores_gemma":[0.0009203627,0.0042195446,0.889527,0.0028503458,0.00071425544,0.0021867757,0.06974929,0.004840371,0.0012565076,0.0009376393,0.02258638,0.00021138525],"about_ca_topic_score_codex":0.006499951,"about_ca_topic_score_gemma":0.00610774,"teacher_disagreement_score":0.96518075,"about_ca_system_score_codex":0.003269228,"about_ca_system_score_gemma":0.0055572693,"threshold_uncertainty_score":0.1841439},"labels":[],"label_agreement":null},{"id":"W3009958360","doi":"10.2196/16948","title":"Semantic Deep Learning: Prior Knowledge and a Type of Four-Term Embedding Analogy to Acquire Treatments for Well-Known Diseases","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Analogy; Pairwise comparison; Artificial intelligence; Natural language processing; Computer science; Embedding; Deep learning; Term (time); Information retrieval; Linguistics","score_opus":0.027762147940074082,"score_gpt":0.34129162653837997,"score_spread":0.3135294785983059,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3009958360","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21469022,0.0068433345,0.73204297,0.006887893,0.00038473093,0.0006565542,0.014956634,0.003042944,0.020494588],"genre_scores_gemma":[0.6318789,0.0011393658,0.34436017,0.0009008391,0.00012617145,0.000580296,0.017456224,0.00010092763,0.0034570554],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99833834,0.00077562284,0.00017693544,0.00045906036,0.00018696878,0.00006305838],"domain_scores_gemma":[0.995129,0.0035719513,0.00027902078,0.0006263334,0.00030271307,0.00009093767],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021023313,0.0008080343,0.0003663883,0.0022917981,0.00048439292,0.0010792825,0.00093858544,0.0010628434,0.0051187435],"category_scores_gemma":[0.011815576,0.00025678868,0.0010067654,0.0018916838,0.0008428266,0.0037774928,0.0017534439,0.0018276166,0.0010980034],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006924249,0.00065835397,0.021198362,0.0036167812,0.0005093559,0.0008560758,0.0020741874,0.029450476,0.014048777,0.073406704,0.03070479,0.8227837],"study_design_scores_gemma":[0.00023908605,0.0006182772,0.020527918,0.0016171719,0.00056211144,0.0010377519,0.0021037858,0.45078948,0.022059532,0.4013069,0.098975025,0.00016297148],"about_ca_topic_score_codex":0.0021036512,"about_ca_topic_score_gemma":0.004650567,"teacher_disagreement_score":0.0051187435,"about_ca_system_score_codex":0.0009820408,"about_ca_system_score_gemma":0.0015201712,"threshold_uncertainty_score":0.017123938},"labels":[],"label_agreement":null},{"id":"W3011478582","doi":"10.2196/17643","title":"A Graph Convolutional Network–Based Method for Chemical-Protein Interaction Extraction: Algorithm Development","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Key Research and Development Program of China","keywords":"Computer science; Sentence; Dependency graph; ENCODE; Natural language processing; Artificial intelligence; Graph; Information extraction; Relationship extraction; Dependency (UML); Knowledge graph; Biomedical text mining; Machine learning; Theoretical computer science; Text mining","score_opus":0.0294580438385864,"score_gpt":0.3368170228323837,"score_spread":0.30735897899379727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3011478582","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010412021,0.0007265989,0.98199445,0.0004669235,0.00007552829,0.00020290713,0.000303422,0.0039318,0.0018863191],"genre_scores_gemma":[0.18349636,0.0013841826,0.8008634,0.00051153696,0.000100528465,0.0005934264,0.0026738527,0.00039108624,0.009985714],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99975127,0.000031906824,0.000018906421,0.000082833125,0.000077293094,0.000037721195],"domain_scores_gemma":[0.9996737,0.000110199675,0.00002892887,0.000033616052,0.00013232033,0.000021213677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006605428,0.0012274846,0.00068041007,0.0013567983,0.0005173205,0.00061910244,0.0018111418,0.0012176632,0.003985833],"category_scores_gemma":[0.0013218512,0.00054427265,0.0008486567,0.001280073,0.0004614189,0.0012872546,0.0008965169,0.0013957841,0.001474468],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013222829,0.00015726671,0.0016220165,0.00021903834,0.00015613053,0.00022108747,0.000055192384,0.29511398,0.014112881,0.0077478555,0.010552197,0.6699101],"study_design_scores_gemma":[0.000008286724,0.000012290463,0.00016082483,0.0000064506366,0.000014514349,0.000033289834,0.0000045081215,0.99469405,0.0023574687,0.0014396495,0.0012630813,0.0000056826325],"about_ca_topic_score_codex":0.030167716,"about_ca_topic_score_gemma":0.035976715,"teacher_disagreement_score":0.030167716,"about_ca_system_score_codex":0.0016066786,"about_ca_system_score_gemma":0.002192321,"threshold_uncertainty_score":0.059984267},"labels":[],"label_agreement":null},{"id":"W3012108798","doi":"10.2196/17644","title":"Document-Level Biomedical Relation Extraction Leveraging Pretrained Self-Attention Structure and Entity Replacement: Algorithm and Pretreatment Method Validation Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"Relationship extraction; Computer science; Preprocessor; Relation (database); Semantics (computer science); Sentence; Context (archaeology); Noise (video); Artificial intelligence; Natural language processing; Heuristic; Data mining; Biomedical text mining; Information retrieval; Machine learning; Pattern recognition (psychology); Text mining; Image (mathematics)","score_opus":0.021014738332400675,"score_gpt":0.3257688454509995,"score_spread":0.3047541071185988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3012108798","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24911948,0.0089197885,0.6992852,0.0010382191,0.0006006365,0.0015763029,0.0066881385,0.025775263,0.0069969203],"genre_scores_gemma":[0.34176672,0.00193615,0.6163245,0.00044178014,0.00012727315,0.0008759431,0.031057382,0.00042946887,0.00704074],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983047,0.00032775378,0.0002125879,0.0006970645,0.00033636804,0.00012154319],"domain_scores_gemma":[0.99603444,0.0015183813,0.00024938822,0.00096223544,0.0011329029,0.00010277812],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002920009,0.001590053,0.0009740721,0.0027571318,0.000777499,0.0011766545,0.0018725288,0.0015364168,0.003047212],"category_scores_gemma":[0.0068702656,0.00039597342,0.0016317074,0.0022799487,0.000504104,0.0023773059,0.0011573667,0.0018303423,0.002469863],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048016434,0.00050233287,0.009991102,0.00074403005,0.00036669552,0.0004746114,0.00029517958,0.027406208,0.02715041,0.0019073689,0.019326156,0.9113558],"study_design_scores_gemma":[0.00013985511,0.0005152237,0.016924184,0.00017351288,0.0003462841,0.001072033,0.00038082368,0.86697614,0.08477473,0.0030861625,0.025524873,0.00008615443],"about_ca_topic_score_codex":0.0102393525,"about_ca_topic_score_gemma":0.012547916,"teacher_disagreement_score":0.0102393525,"about_ca_system_score_codex":0.0011700495,"about_ca_system_score_gemma":0.0021058135,"threshold_uncertainty_score":0.020359516},"labels":[],"label_agreement":null},{"id":"W3014082802","doi":"10.2196/12799","title":"Identification of the Best Semantic Expansion to Query PubMed Through Automatic Performance Assessment of Four Search Strategies on All Medical Subject Heading Descriptors: Comparative Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Unified Medical Language System; Computer science; Information retrieval; Synonym (taxonomy); Precision and recall; Search engine indexing; Subject (documents); Construct (python library); Term (time); Query expansion; Semantics (computer science); Measure (data warehouse); Index (typography); Data mining; World Wide Web","score_opus":0.09308255225453313,"score_gpt":0.38443317167536334,"score_spread":0.2913506194208302,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014082802","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.846825,0.016178602,0.1143116,0.00084752205,0.00020675777,0.0017936669,0.0053345915,0.0075939735,0.0069082654],"genre_scores_gemma":[0.80864334,0.0021118128,0.1806917,0.00023610536,0.000104947634,0.0011222778,0.006038924,0.0003882887,0.00066252064],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9901467,0.00490977,0.0016622787,0.001512835,0.0015431277,0.00022528291],"domain_scores_gemma":[0.94509083,0.045150332,0.0025747658,0.0014241845,0.005350883,0.00040903912],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.015205784,0.001574978,0.0015496144,0.012051253,0.00055466645,0.0022588335,0.0010124443,0.0011601113,0.0013108517],"category_scores_gemma":[0.06276522,0.00027313616,0.0016622886,0.0060568186,0.0005148516,0.0027122626,0.0011164922,0.0004178295,0.0007125196],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008317553,0.0011325764,0.10917821,0.011190485,0.002831526,0.00054305757,0.0019243334,0.014277273,0.04720669,0.0022627055,0.010394711,0.7907409],"study_design_scores_gemma":[0.0025505046,0.014658669,0.3030566,0.0016602188,0.012803181,0.0042826594,0.004266203,0.4919135,0.13074875,0.008830875,0.024561124,0.0006676825],"about_ca_topic_score_codex":0.0029672806,"about_ca_topic_score_gemma":0.0036225098,"teacher_disagreement_score":0.9847942,"about_ca_system_score_codex":0.0011172832,"about_ca_system_score_gemma":0.0022638792,"threshold_uncertainty_score":0.08041686},"labels":[],"label_agreement":null},{"id":"W3014334937","doi":"10.36227/techrxiv.12055998.v1","title":"Drug-Drug Interaction Detection (DDI) Over the Social Media using Convolutional Neural Networks","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Drug; Convolutional neural network; Computer science; Drug-drug interaction; Artificial intelligence; Pharmacology; Medicine","score_opus":0.03819318257508286,"score_gpt":0.30870577643930425,"score_spread":0.2705125938642214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014334937","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4632483,0.0124874385,0.4028049,0.005314663,0.000549591,0.00088914647,0.07432382,0.023225812,0.01715636],"genre_scores_gemma":[0.7357453,0.004080879,0.18897742,0.00047528086,0.00022474416,0.00032047028,0.058068413,0.00037319277,0.011734295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993431,0.00012499162,0.000058054593,0.00018032778,0.00022006531,0.0000734715],"domain_scores_gemma":[0.9985934,0.0007653494,0.00018528514,0.00019059257,0.00021219032,0.000053324165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009816006,0.00095882325,0.00048434,0.0032829237,0.00038957025,0.00085339166,0.00045026708,0.00067822216,0.0014795158],"category_scores_gemma":[0.002945391,0.00028373697,0.00082724745,0.0021438643,0.00024446368,0.0015431613,0.00085657253,0.00066268607,0.00090470666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094850274,0.000765246,0.032606855,0.0013297753,0.001015312,0.00091432524,0.00029176945,0.0487745,0.04518706,0.007828127,0.04284672,0.8174918],"study_design_scores_gemma":[0.00007944658,0.0001822959,0.021549094,0.00016961133,0.0003787478,0.0004093479,0.0001349895,0.8916871,0.039327756,0.017378913,0.028650686,0.00005199715],"about_ca_topic_score_codex":0.0117440475,"about_ca_topic_score_gemma":0.020585146,"teacher_disagreement_score":0.0117440475,"about_ca_system_score_codex":0.0009227976,"about_ca_system_score_gemma":0.0012614201,"threshold_uncertainty_score":0.023351371},"labels":[],"label_agreement":null},{"id":"W3014437799","doi":"","title":"Canadian EdGEO National Workshop Program","year":2009,"lang":"en","type":"article","venue":"AGUSM","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Library science; Computer science","score_opus":0.016592472918983088,"score_gpt":0.30596019155703885,"score_spread":0.28936771863805577,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3014437799","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0043601594,0.0030038925,0.0042977836,0.046373274,0.0068703284,0.00063401717,0.040090594,0.0017967005,0.8925733],"genre_scores_gemma":[0.007994985,0.0014625372,0.0053870333,0.0018138805,0.0002694327,0.00013904752,0.013667702,0.0003772217,0.9688882],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9971324,0.00015671899,0.000042717173,0.0002498387,0.0018253322,0.0005929694],"domain_scores_gemma":[0.99334174,0.00020560325,0.0000720518,0.0002467926,0.0036403479,0.0024935587],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038521946,0.0010059865,0.00055290904,0.0038816142,0.004954211,0.005931663,0.00263997,0.0016957887,0.2669732],"category_scores_gemma":[0.003305434,0.00042444112,0.0009582644,0.0029655134,0.0009306667,0.0014823985,0.0036940796,0.0013427774,0.07694062],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000057067442,0.000038299608,0.00032042892,0.000050930325,0.0000030994956,0.00004664808,0.00008084817,0.00011458417,0.00029917547,0.0071066506,0.9615602,0.03032207],"study_design_scores_gemma":[0.000008820983,0.000003905804,0.000851092,0.000035221732,0.0000025804245,0.000014351648,0.00015203336,0.00015745504,0.00015173368,0.00074674824,0.99786985,0.000006192951],"about_ca_topic_score_codex":0.8163834,"about_ca_topic_score_gemma":0.9273727,"teacher_disagreement_score":0.2669732,"about_ca_system_score_codex":0.023719221,"about_ca_system_score_gemma":0.0967555,"threshold_uncertainty_score":0.8931143},"labels":[],"label_agreement":null},{"id":"W3017538732","doi":"10.2196/17638","title":"Document-Level Biomedical Relation Extraction Using Graph Convolutional Network and Multihead Attention: Algorithm Development and Validation","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Natural Science Foundation of China","keywords":"Computer science; Relationship extraction; Artificial intelligence; Word embedding; Dependency graph; Graph; Natural language processing; Information extraction; Precision and recall; Convolutional neural network; Dependency (UML); Context (archaeology); Embedding; Theoretical computer science","score_opus":0.040842741485564,"score_gpt":0.31475972683180403,"score_spread":0.27391698534624004,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017538732","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.31925058,0.006170419,0.63734215,0.0011563423,0.0003553856,0.0012834446,0.0022611173,0.02623957,0.005940896],"genre_scores_gemma":[0.57162285,0.001393439,0.41213077,0.0005491211,0.00007260583,0.0007711301,0.0076126214,0.00037363992,0.0054738526],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991954,0.00015207948,0.000064620275,0.000270239,0.00018702153,0.00013064739],"domain_scores_gemma":[0.9978967,0.0009627688,0.00015044998,0.00026489186,0.00065371196,0.00007143693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022057856,0.0020167092,0.0010993714,0.002217816,0.00070301106,0.0010142685,0.0025745707,0.0021783456,0.002470225],"category_scores_gemma":[0.004369531,0.00065207965,0.0012160913,0.001600044,0.00060703186,0.0017995781,0.001166313,0.0019470396,0.00095263],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004925505,0.0006210197,0.0069910544,0.00041407172,0.00039545324,0.00035483783,0.000087900415,0.2689268,0.014017118,0.0017288534,0.008582491,0.6973878],"study_design_scores_gemma":[0.00003081662,0.00005399894,0.00075789186,0.000011774891,0.00003778898,0.00004112656,0.000013819425,0.99362,0.0043193386,0.00062964513,0.00047512204,0.000008761253],"about_ca_topic_score_codex":0.04781404,"about_ca_topic_score_gemma":0.046675604,"teacher_disagreement_score":0.04781404,"about_ca_system_score_codex":0.00288258,"about_ca_system_score_gemma":0.0032703623,"threshold_uncertainty_score":0.095071495},"labels":[],"label_agreement":null},{"id":"W3017715546","doi":"10.1162/dint_a_00058","title":"The Semantic Data Dictionary – An Approach for Describing and Annotating Data","year":2020,"lang":"en","type":"article","venue":"Data Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":49,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"CARE Canada; Privacy Analytics (Canada)","funders":"National Institute of Environmental Health Sciences","keywords":"Computer science; Natural language processing; Data dictionary; Information retrieval; Artificial intelligence; World Wide Web; Metadata","score_opus":0.3725294609797752,"score_gpt":0.37406425153345557,"score_spread":0.0015347905536803874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3017715546","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001125762,0.00050260714,0.98273164,0.002434977,0.0003308161,0.0007745255,0.0046403883,0.001840637,0.0056187618],"genre_scores_gemma":[0.013279613,0.0009341579,0.96846306,0.001407623,0.00014898027,0.0013282051,0.011383435,0.0009005755,0.002154369],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9608474,0.01607397,0.010315718,0.0047968104,0.0071904142,0.0007756839],"domain_scores_gemma":[0.92899406,0.02610219,0.0049694953,0.027786447,0.010478642,0.0016692375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034431316,0.0018237968,0.0024186429,0.020726489,0.005183016,0.013950745,0.006920692,0.0041535487,0.005078309],"category_scores_gemma":[0.06554477,0.0022999144,0.0043082377,0.024620656,0.008127064,0.028139997,0.0157383,0.009992349,0.0041063046],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016435627,0.00013103274,0.0023987575,0.0018497347,0.00023705387,0.00059153925,0.008423929,0.0047395816,0.004845076,0.78626823,0.040177904,0.15017271],"study_design_scores_gemma":[0.000041396765,0.000062125306,0.00070492283,0.0016058718,0.00011163087,0.00076965534,0.0033589923,0.013745306,0.004673659,0.32289666,0.6518505,0.00017921776],"about_ca_topic_score_codex":0.013728149,"about_ca_topic_score_gemma":0.013377624,"teacher_disagreement_score":0.034431316,"about_ca_system_score_codex":0.004709535,"about_ca_system_score_gemma":0.016220227,"threshold_uncertainty_score":0.18209243},"labels":[],"label_agreement":null},{"id":"W3020901151","doi":"10.1002/9781118445112.stat06266","title":"Graphical Presentation of Longitudinal Data","year":2014,"lang":"en","type":"other","venue":"Wiley StatsRef: Statistics Reference Online","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; TRACE (psycholinguistics); Graphics; Presentation (obstetrics); Representation (politics); Simple (philosophy); Process (computing); Longitudinal data; Theoretical computer science; Statistical graphics; Data science; Computer graphics (images); Programming language; Linguistics; Data mining; Epistemology","score_opus":0.06869424596413887,"score_gpt":0.3631085233084239,"score_spread":0.29441427734428505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3020901151","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065997634,0.0047822623,0.43007806,0.031007672,0.012840547,0.0011387405,0.29604828,0.10329799,0.11420669],"genre_scores_gemma":[0.14867386,0.009426242,0.4671777,0.015192066,0.005864998,0.004048338,0.18164812,0.030297842,0.1376709],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9949091,0.0029499128,0.00048397554,0.0004720835,0.0010270078,0.00015788962],"domain_scores_gemma":[0.93973947,0.039962143,0.0041138595,0.004967645,0.010122317,0.0010945878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0066262316,0.0016714472,0.0008962226,0.007627238,0.00065516785,0.0041884664,0.0016934486,0.0013307368,0.26638103],"category_scores_gemma":[0.0714819,0.0005343441,0.0011212325,0.005533942,0.0006553649,0.0028297203,0.002812455,0.0018380905,0.052139986],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027207012,0.000034675508,0.0013471049,0.0012462189,0.00009422006,0.00026297913,0.0005192165,0.0025915732,0.00097482756,0.03152066,0.84806037,0.113076046],"study_design_scores_gemma":[0.00009696501,0.000046511675,0.001350056,0.0008481651,0.00005801592,0.00030585701,0.0002036633,0.0057632853,0.0013412785,0.03899066,0.9509092,0.00008619766],"about_ca_topic_score_codex":0.0024209805,"about_ca_topic_score_gemma":0.0026127307,"teacher_disagreement_score":0.26638103,"about_ca_system_score_codex":0.0013213702,"about_ca_system_score_gemma":0.0019714395,"threshold_uncertainty_score":0.8911333},"labels":[],"label_agreement":null},{"id":"W3022705289","doi":"","title":"LibGuides: Vancouver Referencing Style @ UCT: Resources used","year":2014,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Computer science; Artificial intelligence; Geography; Archaeology","score_opus":0.02217327219357417,"score_gpt":0.26642183765649224,"score_spread":0.24424856546291807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3022705289","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006388958,0.0019088234,0.050746385,0.0025252437,0.0015411337,0.0006528932,0.37531617,0.1109342,0.45573628],"genre_scores_gemma":[0.004568396,0.003279504,0.060761,0.0015972976,0.00037195758,0.0013422866,0.49437055,0.1457335,0.28797558],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965006,0.0005624091,0.00048841705,0.00047614836,0.0017134092,0.00025887685],"domain_scores_gemma":[0.98678464,0.002593542,0.00041529094,0.0022662885,0.00713877,0.0008015771],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0031781844,0.0025335322,0.002203629,0.013843147,0.0034145454,0.010945212,0.005353194,0.0021421541,0.5048142],"category_scores_gemma":[0.023692206,0.0021765088,0.0016237383,0.025871428,0.0016020614,0.009147537,0.008386086,0.0039246087,0.5187569],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020189476,0.000006624576,0.000057787936,0.0003770666,0.0000069567322,0.000029106692,0.0001397305,0.00005997699,0.0001919963,0.003157908,0.97256887,0.023383804],"study_design_scores_gemma":[0.0000059739787,0.0000018679264,0.00011471178,0.00017798926,0.000005236556,0.000047565234,0.00006356258,0.00011086429,0.0004470798,0.0016833639,0.99732363,0.00001798046],"about_ca_topic_score_codex":0.09487728,"about_ca_topic_score_gemma":0.13086194,"teacher_disagreement_score":0.9890548,"about_ca_system_score_codex":0.005021231,"about_ca_system_score_gemma":0.010583628,"threshold_uncertainty_score":0.70632243},"labels":[],"label_agreement":null},{"id":"W3023407153","doi":"10.1007/978-981-15-1412-8_9","title":"Organization of Breast Imaging Reports","year":2020,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Mount Sinai Hospital; Women's College Hospital","funders":"","keywords":"BI-RADS; Breast cancer; Breast imaging; Medical physics; Medicine; Computer science; Cancer; Mammography; Internal medicine","score_opus":0.0075493499987798595,"score_gpt":0.21412325991149372,"score_spread":0.20657390991271388,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3023407153","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02009605,0.012714536,0.45826727,0.0050755218,0.0013546649,0.0012338783,0.13325143,0.016301895,0.35170466],"genre_scores_gemma":[0.06480531,0.013282184,0.50922894,0.0011295113,0.0007676733,0.00054402684,0.20337087,0.004117023,0.20275447],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.999428,0.000092574606,0.00008545396,0.00014575895,0.00020104094,0.000047284873],"domain_scores_gemma":[0.9987645,0.00037203525,0.00016776317,0.00025123064,0.00037511266,0.000069283604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007406429,0.0006548479,0.0004337739,0.0075629978,0.0006436347,0.004072483,0.00082138786,0.0004947197,0.028588902],"category_scores_gemma":[0.002687993,0.0004321132,0.00071424345,0.0077977423,0.00045775884,0.002528736,0.0012944624,0.00055124384,0.025428092],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000083335195,0.000053340733,0.0019937432,0.0007388633,0.000027676977,0.00042110228,0.00044252473,0.0023736085,0.010832617,0.056542482,0.1661082,0.7603825],"study_design_scores_gemma":[0.000009853788,0.000027242178,0.0038606066,0.0003516765,0.000047775735,0.0011386173,0.00038571708,0.004813679,0.0088508865,0.0308725,0.94961506,0.000026499558],"about_ca_topic_score_codex":0.0035087147,"about_ca_topic_score_gemma":0.0037879273,"teacher_disagreement_score":0.028588902,"about_ca_system_score_codex":0.0008838936,"about_ca_system_score_gemma":0.0014444718,"threshold_uncertainty_score":0.09563935},"labels":[],"label_agreement":null},{"id":"W3023815642","doi":"10.1007/978-3-030-47358-7_34","title":"Detection and Diagnosis of Breast Cancer Using a Bayesian Approach","year":2020,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Bayesian network; Breast cancer; Bayesian probability; Computer science; Conditional probability; Classifier (UML); Artificial intelligence; Predictive value; Sensitivity (control systems); Machine learning; Pattern recognition (psychology); Cancer; Algorithm; Statistics; Mathematics; Medicine; Internal medicine","score_opus":0.018992236153140267,"score_gpt":0.25650052346431723,"score_spread":0.23750828731117696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3023815642","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011697895,0.0016945518,0.9831405,0.00057910656,0.000042182117,0.00006652006,0.00026691973,0.00040887957,0.0021034924],"genre_scores_gemma":[0.37579906,0.0025954742,0.6139192,0.00068844814,0.00032816987,0.0002990602,0.0011592102,0.00010874947,0.005102695],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998664,0.00035585507,0.00011496195,0.00028569894,0.00047905656,0.000100377496],"domain_scores_gemma":[0.9972427,0.0021061278,0.00013315289,0.000099566314,0.00036746537,0.000051058083],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022150509,0.0007419929,0.0013139128,0.002147602,0.0005451779,0.0013012678,0.0018809161,0.0016023727,0.0022515163],"category_scores_gemma":[0.0070698815,0.00091443205,0.0016281302,0.0009792994,0.00046655646,0.0014834132,0.0011446946,0.0011013341,0.0011770872],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000616166,0.00046435656,0.010201462,0.00049051724,0.0005278247,0.0003990401,0.00022638272,0.1811029,0.016428875,0.015940815,0.006700808,0.76690096],"study_design_scores_gemma":[0.000067174136,0.00014376489,0.003503924,0.00010804905,0.00027126612,0.0005944203,0.00005146119,0.9440066,0.003826971,0.04434723,0.0030192798,0.00005981079],"about_ca_topic_score_codex":0.006319345,"about_ca_topic_score_gemma":0.008833195,"teacher_disagreement_score":0.006319345,"about_ca_system_score_codex":0.000809278,"about_ca_system_score_gemma":0.0012812921,"threshold_uncertainty_score":0.012565136},"labels":[],"label_agreement":null},{"id":"W3031309641","doi":"","title":"Temporal Histories of Epidemic Events (THEE): A Case Study in Temporal Annotation for Public Health","year":2020,"lang":"en","type":"article","venue":"Language Resources and Evaluation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Agency of Canada; University of Toronto","funders":"","keywords":"Annotation; Computer science; Metadata; Domain (mathematical analysis); Event (particle physics); Public domain; Information retrieval; Process (computing); Temporal annotation; Style (visual arts); Natural language processing; Artificial intelligence; World Wide Web; History; Natural language","score_opus":0.10611195314227607,"score_gpt":0.38414537583421693,"score_spread":0.27803342269194087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3031309641","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.45811802,0.004313724,0.35909176,0.023391493,0.0005723654,0.0018035175,0.08893823,0.005998227,0.057772692],"genre_scores_gemma":[0.69975775,0.0017769837,0.25279933,0.001084466,0.00013041451,0.0006450691,0.034447096,0.0008909113,0.008468001],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9950819,0.0023980301,0.00069786457,0.0005816004,0.0010723592,0.00016823814],"domain_scores_gemma":[0.9345662,0.05384064,0.0034624748,0.002972795,0.004094421,0.0010635541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011208732,0.00041004037,0.00031335675,0.0040950803,0.001968867,0.0022609774,0.0011847443,0.0012828966,0.0037511042],"category_scores_gemma":[0.035589248,0.00025812525,0.0005320865,0.006504504,0.0013163966,0.0053196033,0.0024115993,0.0012959995,0.0006549497],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019250576,0.0007954218,0.15611261,0.0064433375,0.00026961314,0.014136548,0.054091778,0.023577573,0.015136072,0.21556996,0.08882712,0.4231149],"study_design_scores_gemma":[0.0001421909,0.0002914542,0.07159671,0.0020517183,0.00036825455,0.007154112,0.03912747,0.12531523,0.022558775,0.09470068,0.6364335,0.00025989904],"about_ca_topic_score_codex":0.038677078,"about_ca_topic_score_gemma":0.046086323,"teacher_disagreement_score":0.038677078,"about_ca_system_score_codex":0.002408865,"about_ca_system_score_gemma":0.004569238,"threshold_uncertainty_score":0.07690388},"labels":[],"label_agreement":null},{"id":"W3031501022","doi":"","title":"Automated Analysis of Public Health Laboratory Test Results.","year":2020,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Interpretability; Generalizability theory; Computer science; Artificial intelligence; Machine learning; Class (philosophy); Test (biology); Code (set theory); Natural language processing; Data mining; Programming language; Statistics; Mathematics","score_opus":0.056054776079371274,"score_gpt":0.2789610636699819,"score_spread":0.22290628759061065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3031501022","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17737862,0.017752811,0.51271707,0.0068949084,0.0013319555,0.0024616309,0.2114828,0.032951627,0.037028525],"genre_scores_gemma":[0.48779365,0.0053730807,0.36438936,0.0011103894,0.0006135837,0.0008276081,0.13375029,0.0006948781,0.0054472033],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9926564,0.0021502539,0.0011610513,0.0010940759,0.0027485872,0.00018961362],"domain_scores_gemma":[0.9543436,0.021544963,0.008334941,0.0047959224,0.010482483,0.00049803586],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0060948865,0.00084200146,0.00069390825,0.012774224,0.00055415416,0.0024591987,0.001064343,0.0005921752,0.0036839445],"category_scores_gemma":[0.03401871,0.00021778558,0.00072964345,0.006657399,0.00043412263,0.0018288376,0.0013086394,0.00061349705,0.0034649],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068373,0.00039198998,0.09396329,0.0039336365,0.00037937978,0.0011146795,0.0011020793,0.0045542093,0.019541694,0.005830508,0.06269974,0.8058051],"study_design_scores_gemma":[0.00022604266,0.00078622566,0.2751413,0.0033186586,0.0010761004,0.0052963654,0.00455463,0.13420966,0.12416614,0.06062086,0.39029106,0.00031292514],"about_ca_topic_score_codex":0.0030474407,"about_ca_topic_score_gemma":0.003269195,"teacher_disagreement_score":0.012774224,"about_ca_system_score_codex":0.0008185878,"about_ca_system_score_gemma":0.0025839428,"threshold_uncertainty_score":0.03223324},"labels":[],"label_agreement":null},{"id":"W3034095843","doi":"10.14745/ccdr.46i06a06","title":"Application of natural language processing algorithms for extracting information from news articles in event-based surveillance.","year":2020,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Health Canada; Response Biomedical (Canada); University of Toronto; Public Health Agency of Canada","funders":"","keywords":"Information extraction; Computer science; Event (particle physics); Focus (optics); Named-entity recognition; Information retrieval; Data science; Artificial intelligence; Natural language processing; Task (project management); Engineering","score_opus":0.018948306375277342,"score_gpt":0.2632110314590011,"score_spread":0.24426272508372376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034095843","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040560916,0.013884905,0.85371935,0.0064793485,0.0011080537,0.0039274083,0.045525417,0.013819718,0.020974811],"genre_scores_gemma":[0.07877881,0.0046932055,0.87911814,0.000624579,0.00032212658,0.0008734441,0.03206458,0.00036037818,0.0031647736],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99558085,0.0015668113,0.0010050335,0.00064842263,0.0010955258,0.000103323975],"domain_scores_gemma":[0.9738287,0.020145297,0.001864624,0.0008838649,0.0030395538,0.00023800383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006204484,0.0011265074,0.00076950673,0.012518561,0.0010926594,0.004012974,0.00092645193,0.0009874216,0.003821266],"category_scores_gemma":[0.022436885,0.00044518468,0.00131256,0.007577488,0.0007161378,0.0031860815,0.0014671729,0.0013503206,0.0032961487],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036157944,0.00041647913,0.010353588,0.008677178,0.0005930148,0.0019519387,0.002211924,0.004290393,0.039965726,0.015583451,0.046895657,0.86869913],"study_design_scores_gemma":[0.00026273192,0.0004915322,0.03522321,0.003756184,0.0015684444,0.0044259313,0.005541547,0.175069,0.12950647,0.09062702,0.5531722,0.0003557384],"about_ca_topic_score_codex":0.003483004,"about_ca_topic_score_gemma":0.005143807,"teacher_disagreement_score":0.012518561,"about_ca_system_score_codex":0.0010951823,"about_ca_system_score_gemma":0.002587124,"threshold_uncertainty_score":0.032812834},"labels":[],"label_agreement":null},{"id":"W3034575096","doi":"10.2196/21379","title":"Correction: Prioritization of Free-Text Clinical Documents: A Novel Use of a Bayesian Classifier","year":2020,"lang":"en","type":"erratum","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Bayesian probability; Prioritization; Artificial intelligence; Classifier (UML); Natural language processing; Naive Bayes classifier; Information retrieval; Machine learning; Data mining; Support vector machine","score_opus":0.03611017026749362,"score_gpt":0.3373994203241242,"score_spread":0.30128925005663054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3034575096","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00031661094,0.000913174,0.0036831666,0.10723832,0.8784054,0.00006688478,0.004465569,0.0012103342,0.0037005013],"genre_scores_gemma":[0.034384724,0.010192227,0.039280616,0.19865699,0.41214055,0.00068052683,0.009848957,0.0063278005,0.2884876],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9921483,0.0015332166,0.0016613287,0.00095390616,0.00325681,0.0004463566],"domain_scores_gemma":[0.92095816,0.029408729,0.0031794722,0.0042635584,0.04003852,0.0021516443],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006686913,0.002429498,0.0017724818,0.005423825,0.0036159772,0.0045291563,0.004096214,0.0074972915,0.06916495],"category_scores_gemma":[0.15044254,0.0012864541,0.0018325296,0.003566023,0.0031337421,0.0026461214,0.0026007858,0.009895718,0.03435494],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041485146,0.0000051915636,0.00009533048,0.00014392799,0.000018161307,0.00045882928,0.000045528206,0.0000713707,0.00006538129,0.00075738603,0.9913276,0.0069698594],"study_design_scores_gemma":[0.000079544276,0.00002519989,0.00074780005,0.00058909075,0.00006905861,0.0019188013,0.00012545509,0.00092211535,0.0006087737,0.0032196376,0.9916276,0.000066976354],"about_ca_topic_score_codex":0.030175189,"about_ca_topic_score_gemma":0.03530379,"teacher_disagreement_score":0.06916495,"about_ca_system_score_codex":0.0038910524,"about_ca_system_score_gemma":0.007859986,"threshold_uncertainty_score":0.2313798},"labels":[],"label_agreement":null},{"id":"W3035019845","doi":"10.3233/jifs-179869","title":"Improving the identification of confused drug names in Spanish","year":2020,"lang":"en","type":"article","venue":"Journal of Intelligent & Fuzzy Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Identification (biology); Similarity (geometry); Drug; Computer science; Medicine; Linguistics; Artificial intelligence; Pharmacology","score_opus":0.021066291212601885,"score_gpt":0.2582285446052753,"score_spread":0.23716225339267338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035019845","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.84362155,0.016109845,0.10090197,0.0034646927,0.0013654316,0.000325092,0.0054523433,0.0048929686,0.02386622],"genre_scores_gemma":[0.95244986,0.0018714458,0.02910662,0.0008185029,0.00023363132,0.00008149305,0.008847253,0.0002981816,0.0062929555],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964168,0.0017861589,0.00029970356,0.0008170059,0.00041518468,0.00026517163],"domain_scores_gemma":[0.9911953,0.00523502,0.000602617,0.0008189238,0.0018844636,0.00026368222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00673607,0.002502668,0.0016826303,0.0029943502,0.0007415985,0.0028743846,0.0018531608,0.0015586532,0.0029785451],"category_scores_gemma":[0.017695598,0.00036606065,0.0023813432,0.0019191235,0.0005623839,0.0035002918,0.0017700401,0.0014522418,0.0017808964],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028832322,0.0011610269,0.16932683,0.001290897,0.0012525663,0.0016922544,0.00093320815,0.26060733,0.003880871,0.0047288765,0.045425374,0.5068175],"study_design_scores_gemma":[0.00016405656,0.00028249226,0.019267455,0.00015490226,0.00039454785,0.00033678245,0.00088095746,0.9564957,0.0029280153,0.004800727,0.014176718,0.00011757962],"about_ca_topic_score_codex":0.04914489,"about_ca_topic_score_gemma":0.025928926,"teacher_disagreement_score":0.04914489,"about_ca_system_score_codex":0.0019272309,"about_ca_system_score_gemma":0.0025665811,"threshold_uncertainty_score":0.0977177},"labels":[],"label_agreement":null},{"id":"W3036470112","doi":"10.3233/shti200283","title":"Intelligent Tools for Precision Public Health","year":2020,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Ontology; Transferability; Public health interventions; Psychological intervention; Computer science; Public health; Set (abstract data type); Population; Software; Data science; Knowledge management; Medicine; Machine learning; Environmental health; Nursing","score_opus":0.18244555400644824,"score_gpt":0.4093353939900585,"score_spread":0.2268898399836103,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3036470112","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011618289,0.0032589224,0.961192,0.009629898,0.0004177828,0.00032791554,0.0017860674,0.008941062,0.013284456],"genre_scores_gemma":[0.029383505,0.002766841,0.9582667,0.00206662,0.0004323179,0.00064835255,0.0033233708,0.00073627755,0.0023759727],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97390443,0.013991439,0.0033901576,0.0031395964,0.0050946684,0.00047968526],"domain_scores_gemma":[0.9141553,0.055768877,0.0041843355,0.020288978,0.004667687,0.0009348092],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03264339,0.0025428364,0.0023648916,0.014882212,0.002120753,0.015157425,0.0044730045,0.004744369,0.011488395],"category_scores_gemma":[0.07901023,0.0016832487,0.0040156352,0.010788557,0.0063940375,0.019412557,0.012599124,0.005835041,0.0056230514],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011592041,0.00013280104,0.0012677445,0.0020308066,0.0004139027,0.00045828833,0.0019380145,0.011443592,0.0015547432,0.7443215,0.025236936,0.2110857],"study_design_scores_gemma":[0.000048224767,0.00002508,0.00027543315,0.0006706027,0.00008327988,0.0001585483,0.00023169731,0.016419796,0.0016119803,0.8169646,0.16345513,0.00005553792],"about_ca_topic_score_codex":0.0033210814,"about_ca_topic_score_gemma":0.0021008924,"teacher_disagreement_score":0.03264339,"about_ca_system_score_codex":0.0031599486,"about_ca_system_score_gemma":0.005798253,"threshold_uncertainty_score":0.1726368},"labels":[],"label_agreement":null},{"id":"W3036480066","doi":"10.3233/shti200128","title":"Clinical Abbreviation Disambiguation Using Deep Contextualized Representation","year":2020,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Artificial intelligence; Representation (politics); Dimension (graph theory); Principal component analysis; Cluster (spacecraft); Word-sense disambiguation; Annotation; Linguistics; Mathematics","score_opus":0.19550265233361913,"score_gpt":0.48390092223796544,"score_spread":0.2883982699043463,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3036480066","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03959525,0.0019415703,0.9361664,0.0010456342,0.0005228659,0.00023488687,0.004413514,0.011650355,0.004429468],"genre_scores_gemma":[0.40133312,0.00086380396,0.58378583,0.0006503802,0.0002437733,0.00028062443,0.009236858,0.00050251547,0.003103186],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988732,0.00023000066,0.00011933939,0.00046170244,0.00020455259,0.0001112637],"domain_scores_gemma":[0.99906963,0.00024742822,0.0001324735,0.0002110247,0.00029844316,0.00004105157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006860763,0.001022491,0.0006937548,0.0026868954,0.00053294306,0.0011178885,0.001284592,0.00080870214,0.002975856],"category_scores_gemma":[0.0032626865,0.0002630863,0.0011579302,0.0022716862,0.0004774359,0.0018307058,0.0018644106,0.001286311,0.0017408562],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047009203,0.00013605107,0.006341146,0.0005525789,0.0001902,0.0009597183,0.0006006972,0.037502747,0.021463858,0.028389191,0.037175715,0.86621803],"study_design_scores_gemma":[0.0000617399,0.00017263988,0.005126156,0.000325771,0.0002774799,0.0013610515,0.0006070043,0.7696806,0.03335026,0.107338525,0.08154509,0.00015371009],"about_ca_topic_score_codex":0.0058708685,"about_ca_topic_score_gemma":0.008449421,"teacher_disagreement_score":0.0058708685,"about_ca_system_score_codex":0.0008670011,"about_ca_system_score_gemma":0.0017204816,"threshold_uncertainty_score":0.011673391},"labels":[],"label_agreement":null},{"id":"W3036643374","doi":"","title":"Guides: Vancouver citation style (based on Citing Medicine): Introduction","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Citation; History; Library science; Computer science; Archaeology","score_opus":0.03039783704103769,"score_gpt":0.27860063478667896,"score_spread":0.24820279774564127,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3036643374","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021456364,0.0015709058,0.024277339,0.005739426,0.004426743,0.00083691714,0.50062364,0.053087316,0.40729204],"genre_scores_gemma":[0.008204484,0.003118654,0.074471876,0.0010798768,0.0015736048,0.0008016637,0.25688288,0.031812675,0.6220542],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99808156,0.0002499065,0.0003605605,0.00022028192,0.0009827889,0.00010498628],"domain_scores_gemma":[0.9757638,0.0054331287,0.0010529305,0.0014408968,0.014645,0.0016641751],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0021435695,0.0014280458,0.0012649623,0.01643637,0.0018393146,0.007684091,0.0020173774,0.0013970535,0.41317332],"category_scores_gemma":[0.02319392,0.0010076253,0.0006242207,0.026499053,0.00063242676,0.0036280027,0.0020360907,0.0015337558,0.28819817],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015109386,0.000010216427,0.00013672949,0.00026762264,0.0000035383814,0.000011897851,0.00005542048,0.000061018716,0.00008464996,0.001070486,0.97332245,0.024960866],"study_design_scores_gemma":[0.000014972829,0.0000066271173,0.0009840961,0.00018684017,0.000007048847,0.00005059435,0.00006513819,0.00030860666,0.00030617177,0.0016739495,0.9963766,0.000019412662],"about_ca_topic_score_codex":0.050240282,"about_ca_topic_score_gemma":0.15478073,"teacher_disagreement_score":0.9923159,"about_ca_system_score_codex":0.0022007695,"about_ca_system_score_gemma":0.0063444655,"threshold_uncertainty_score":0.837037},"labels":[],"label_agreement":null},{"id":"W3036813586","doi":"","title":"Guides: Vancouver citation style (based on Citing Medicine): JBI Connect","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Citation; Style (visual arts); History; Library science; Computer science; Archaeology","score_opus":0.03576146337450676,"score_gpt":0.29009686938258317,"score_spread":0.2543354060080764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3036813586","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011071349,0.0011117458,0.008783251,0.0039367215,0.003840807,0.0005104458,0.39608467,0.036630165,0.54799503],"genre_scores_gemma":[0.0040936098,0.0022454825,0.026414776,0.0007492708,0.0008709293,0.0005812682,0.24052325,0.027171364,0.6973501],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979208,0.00020457244,0.00040567585,0.00032131534,0.0010143361,0.00013322209],"domain_scores_gemma":[0.976448,0.004003086,0.0009459947,0.0018041466,0.014616689,0.0021821538],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0024952318,0.0019313241,0.0019956392,0.021845005,0.0028577514,0.009897405,0.002495075,0.00179764,0.6618533],"category_scores_gemma":[0.029051159,0.0013398424,0.0008373911,0.038217768,0.0007588707,0.0048666876,0.0029387455,0.0023127485,0.5726894],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013557263,0.000008129452,0.00008292326,0.00020062465,0.0000038195403,0.0000089465775,0.000037092792,0.000028819708,0.00006592236,0.000765755,0.9828675,0.015916843],"study_design_scores_gemma":[0.0000165235,0.0000062116783,0.0007678674,0.00021037448,0.000009023607,0.000035289355,0.00006832857,0.00018495995,0.00025979886,0.0014294997,0.99699485,0.000017169863],"about_ca_topic_score_codex":0.05023857,"about_ca_topic_score_gemma":0.146262,"teacher_disagreement_score":0.9901026,"about_ca_system_score_codex":0.0025959725,"about_ca_system_score_gemma":0.008280756,"threshold_uncertainty_score":0.4823252},"labels":[],"label_agreement":null},{"id":"W3037779358","doi":"10.1609/aaai.v34i10.7264","title":"Literature Mining for Incorporating Inductive Bias in Biomedical Prediction Tasks (Student Abstract)","year":2020,"lang":"en","type":"article","venue":"Proceedings of the AAAI Conference on Artificial Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; McGill University","funders":"","keywords":"Leverage (statistics); Inductive bias; Computer science; Machine learning; Relevance (law); Task (project management); Artificial intelligence; Multi-task learning; Feature (linguistics); Training set; Data mining; Data science; Engineering","score_opus":0.1364435877900705,"score_gpt":0.338468309210084,"score_spread":0.2020247214200135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3037779358","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.055241853,0.0015776705,0.930488,0.005204334,0.00035437732,0.00032521412,0.0014534083,0.0028598667,0.0024952304],"genre_scores_gemma":[0.36574617,0.00069563085,0.6258029,0.001606986,0.00079959945,0.00055315613,0.002826811,0.00028103148,0.0016878331],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9917823,0.0041539404,0.0008108036,0.0013124509,0.001728481,0.00021199523],"domain_scores_gemma":[0.90041953,0.07959782,0.005262465,0.0070749614,0.0066237687,0.001021463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02444465,0.0007783745,0.0013648551,0.005980645,0.0012171135,0.0028418663,0.0026713363,0.002034845,0.0027907556],"category_scores_gemma":[0.10978236,0.00075502344,0.0014117379,0.005105735,0.0011727146,0.0043086302,0.0038787944,0.0022580782,0.001600924],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085744815,0.0007853141,0.048228502,0.0012670745,0.0006471438,0.00065788516,0.00069290184,0.05996486,0.008056351,0.031914964,0.029740488,0.817187],"study_design_scores_gemma":[0.00017487897,0.00018090886,0.0058076084,0.00037891764,0.00017999746,0.00050103926,0.00015622261,0.83339405,0.01057869,0.13795683,0.010607613,0.00008323811],"about_ca_topic_score_codex":0.0013076725,"about_ca_topic_score_gemma":0.003326838,"teacher_disagreement_score":0.02444465,"about_ca_system_score_codex":0.0010101801,"about_ca_system_score_gemma":0.0024788454,"threshold_uncertainty_score":0.12927723},"labels":[],"label_agreement":null},{"id":"W3040347603","doi":"","title":"LibGuides: WRHA Virtual Library: Terms of Use","year":2018,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; World Wide Web","score_opus":0.023403031236708283,"score_gpt":0.26542917287041723,"score_spread":0.24202614163370895,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3040347603","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009205028,0.0012099723,0.02119042,0.0015947246,0.0005021556,0.0004664216,0.49206454,0.19350956,0.28854164],"genre_scores_gemma":[0.0077368687,0.0024706125,0.023143424,0.00085443305,0.00028656886,0.0010191482,0.7130503,0.14071931,0.110719375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99825686,0.00023648627,0.00022329915,0.0002480062,0.00078182126,0.00025353278],"domain_scores_gemma":[0.99551743,0.001037156,0.0002590349,0.0012686321,0.0012050528,0.0007126603],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0025145917,0.00289898,0.0021378146,0.010862309,0.0019179372,0.012337157,0.006255619,0.0023714607,0.49238166],"category_scores_gemma":[0.014108393,0.0017546287,0.0017483244,0.017797602,0.0018853543,0.009780079,0.0076459893,0.002737197,0.5721366],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005763559,0.000028115579,0.00013478835,0.0006419682,0.000024313544,0.00008028564,0.00017992368,0.00025241383,0.0003808886,0.0067846784,0.9620396,0.029395407],"study_design_scores_gemma":[0.000020874113,0.0000062079157,0.0003198918,0.00022481995,0.000008266976,0.000095897194,0.00007249567,0.00038150197,0.00048581167,0.0044242246,0.9939281,0.0000319364],"about_ca_topic_score_codex":0.01626103,"about_ca_topic_score_gemma":0.014283598,"teacher_disagreement_score":0.5076183,"about_ca_system_score_codex":0.0035734824,"about_ca_system_score_gemma":0.0065153223,"threshold_uncertainty_score":0.72405595},"labels":[],"label_agreement":null},{"id":"W3040820257","doi":"10.3968/11767","title":"Translator-Author Cooperative Translation Mode Based on Information Theory","year":2020,"lang":"en","type":"article","venue":"Studies in literature and language","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Tact; Translation (biology); Computer science; Mode (computer interface); Process (computing); Field (mathematics); Translation studies; Linguistics; Epistemology; Psychology; Human–computer interaction; Philosophy; Mathematics","score_opus":0.02216793670833665,"score_gpt":0.32422874815931263,"score_spread":0.30206081145097596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3040820257","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00841567,0.00081190944,0.95582294,0.002310709,0.0003505932,0.00042159305,0.00019305461,0.0011561557,0.03051745],"genre_scores_gemma":[0.3205205,0.0014753713,0.65511966,0.0006628005,0.0004806144,0.0011386367,0.0006596293,0.0006102645,0.01933253],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9841929,0.0073276674,0.0018602153,0.002718418,0.0033964526,0.0005043701],"domain_scores_gemma":[0.9799309,0.010587673,0.0013974031,0.0036275892,0.0038709012,0.0005856941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011827781,0.0014401466,0.0011354553,0.006473746,0.0037638005,0.007005185,0.0023866552,0.002449355,0.008667345],"category_scores_gemma":[0.02492444,0.00073408405,0.0017436608,0.0053727822,0.0074102813,0.019221358,0.005509141,0.0021846532,0.004483834],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022405239,0.00010746032,0.0024454554,0.0008324902,0.00011807285,0.0008173423,0.012214603,0.0031502356,0.005186196,0.76044416,0.009386434,0.20507349],"study_design_scores_gemma":[0.00015031418,0.00025369498,0.0011425919,0.00048055875,0.00023701294,0.0015659385,0.0040269033,0.061547063,0.012933724,0.7921516,0.12525854,0.00025199787],"about_ca_topic_score_codex":0.0013449448,"about_ca_topic_score_gemma":0.00094008766,"teacher_disagreement_score":0.011827781,"about_ca_system_score_codex":0.0020793958,"about_ca_system_score_gemma":0.0039720135,"threshold_uncertainty_score":0.062552035},"labels":[],"label_agreement":null},{"id":"W3041913986","doi":"10.17504/protocols.io.bf89jrz6","title":"SOP for populating NCBI submission templates forSARS-CoV-2 (BioSample, SRA, and GenBank) v1","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"GenBank; Metadata; Template; Interoperability; Computer science; Coronavirus disease 2019 (COVID-19); World Wide Web; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Information retrieval; Biology; Programming language; Medicine; Genetics","score_opus":0.0878329148866254,"score_gpt":0.3429870062811122,"score_spread":0.25515409139448675,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3041913986","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018686423,0.00016821934,0.30549008,0.00097584055,0.00059559674,0.0010194611,0.10708598,0.5740475,0.008748617],"genre_scores_gemma":[0.021344457,0.0004923856,0.34237534,0.0015516223,0.00033764986,0.0031155911,0.33975646,0.27907786,0.011948636],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956079,0.0009915621,0.00083364546,0.0007635297,0.001348532,0.00045481956],"domain_scores_gemma":[0.98580605,0.0062013455,0.0007828578,0.004324166,0.0020837162,0.00080180046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011100332,0.003937086,0.0018973888,0.004399618,0.0017240167,0.005804824,0.0036896965,0.003486148,0.103508584],"category_scores_gemma":[0.036212485,0.0033424194,0.0035117655,0.0036251007,0.0011331523,0.005722059,0.008030946,0.0045559257,0.10876696],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015332729,0.00022988777,0.0038215548,0.001549392,0.00029840358,0.0008719724,0.00090901216,0.0026949544,0.011363021,0.019720517,0.859782,0.09722605],"study_design_scores_gemma":[0.0007331379,0.00012427481,0.003270618,0.0006897223,0.00011598056,0.00067466084,0.00047823577,0.026136158,0.028812457,0.041075706,0.89757997,0.0003091547],"about_ca_topic_score_codex":0.0059748474,"about_ca_topic_score_gemma":0.0042313645,"teacher_disagreement_score":0.103508584,"about_ca_system_score_codex":0.0019806195,"about_ca_system_score_gemma":0.0036729714,"threshold_uncertainty_score":0.34627074},"labels":[],"label_agreement":null},{"id":"W3042104358","doi":"10.17504/protocols.io.bhwdj7a6","title":"SARS-CoV-2 EBI submission protocol: ENA, BioSample, and BioProject v1","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Protocol (science); Biology; Metadata; Coronavirus disease 2019 (COVID-19); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); 2019-20 coronavirus outbreak; Computer science; Computational biology; World Wide Web; Virology","score_opus":0.07670131585770588,"score_gpt":0.363572556718384,"score_spread":0.2868712408606781,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3042104358","genre_codex":"dataset","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059828446,0.0017214192,0.220323,0.0055125775,0.0042999336,0.026213042,0.58698076,0.06592911,0.08303725],"genre_scores_gemma":[0.006431015,0.00096117926,0.12551902,0.002423886,0.00057052873,0.032269787,0.7536308,0.025892034,0.052301686],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9855012,0.0044058915,0.0030492155,0.0016430925,0.004115807,0.0012847751],"domain_scores_gemma":[0.97450304,0.0047854246,0.0012148141,0.009239669,0.008821129,0.001436001],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02024731,0.0019937635,0.0023680371,0.0040294398,0.0040118974,0.006523346,0.0046266583,0.0030265527,0.22238539],"category_scores_gemma":[0.038530137,0.0030507618,0.0013107068,0.0032651112,0.001309029,0.0045702485,0.0068073478,0.0052266074,0.34672728],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009563584,0.00016373351,0.00076075253,0.0011955973,0.000058281454,0.0001770118,0.00038430069,0.0001684024,0.01985883,0.0061334316,0.9462308,0.023912506],"study_design_scores_gemma":[0.00018088629,0.000099844605,0.0015378654,0.00046536475,0.000028864346,0.00026123275,0.00025133914,0.00041495136,0.024890553,0.003937433,0.9678371,0.00009466702],"about_ca_topic_score_codex":0.0028722303,"about_ca_topic_score_gemma":0.0034183785,"teacher_disagreement_score":0.22238539,"about_ca_system_score_codex":0.0024080584,"about_ca_system_score_gemma":0.0075782454,"threshold_uncertainty_score":0.7439532},"labels":[],"label_agreement":null},{"id":"W3042438073","doi":"10.2196/17964","title":"Visualization Environment for Federated Knowledge Graphs: Development of an Interactive Biomedical Query Language and Web Application Interface","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; National Institutes of Health","keywords":"Computer science; Query language; Information retrieval; Web search query; Query expansion; Web query classification; Visualization; World Wide Web; Data mining; Search engine","score_opus":0.01368877451368831,"score_gpt":0.32356094656216533,"score_spread":0.309872172048477,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3042438073","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0028193502,0.000096344826,0.9097335,0.00069837103,0.000044284458,0.0005310242,0.0013718974,0.08197801,0.002727259],"genre_scores_gemma":[0.03995055,0.00028813316,0.9397433,0.0009528297,0.0000437749,0.0012212917,0.0054049892,0.008317996,0.004077078],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99639505,0.0010062924,0.00049630576,0.0006377803,0.0012837418,0.00018086137],"domain_scores_gemma":[0.9912129,0.0047829333,0.00043396326,0.0011172858,0.0019452446,0.00050760276],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009010349,0.0015062016,0.0009651511,0.002346336,0.00078431686,0.0044375258,0.0037491962,0.0020567393,0.012799396],"category_scores_gemma":[0.013938483,0.0011628135,0.0021937515,0.0013885312,0.0010818301,0.0054475334,0.004559473,0.0029477684,0.004001186],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021923543,0.0010782087,0.007603373,0.0026259308,0.0005193822,0.0035331396,0.0075719394,0.045267355,0.07817104,0.1598725,0.19297528,0.49858963],"study_design_scores_gemma":[0.000534515,0.00034195554,0.0022173773,0.00067238306,0.00016983865,0.0017980114,0.00065932336,0.51565236,0.077055,0.075768426,0.32472625,0.0004045265],"about_ca_topic_score_codex":0.0045550456,"about_ca_topic_score_gemma":0.0036421588,"teacher_disagreement_score":0.012799396,"about_ca_system_score_codex":0.0017932977,"about_ca_system_score_gemma":0.002328497,"threshold_uncertainty_score":0.047651827},"labels":[],"label_agreement":null},{"id":"W3045040978","doi":"","title":"Guides: Vancouver citation style (based on Citing Medicine): Standards","year":2011,"lang":"en","type":"libguides","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Style (visual arts); Citation; Computer science; Art; World Wide Web; Visual arts","score_opus":0.03344199578688672,"score_gpt":0.3054827839486708,"score_spread":0.27204078816178406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3045040978","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022413642,0.0015860748,0.06342724,0.009369574,0.0038743091,0.0019501091,0.40690157,0.07598463,0.4346652],"genre_scores_gemma":[0.007464745,0.0030521287,0.14426172,0.0016341622,0.00096980494,0.002117127,0.30325872,0.040959667,0.49628195],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9917435,0.0009788986,0.0023694444,0.00062828587,0.00399439,0.0002854675],"domain_scores_gemma":[0.90871274,0.016414544,0.0040957713,0.0070585315,0.060439344,0.003279143],"candidate_categories":["scholarly_communication","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006683984,0.001667441,0.0019184921,0.023236772,0.0031721625,0.012376362,0.0039299177,0.0029851655,0.34227723],"category_scores_gemma":[0.06262933,0.0017693174,0.00090377737,0.04293451,0.0012752782,0.0075037847,0.0031945156,0.0032391408,0.31533206],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027947552,0.000025693584,0.00014802368,0.0005331443,0.00000611387,0.000017613545,0.00013823509,0.000100925485,0.00019753618,0.003234626,0.9554378,0.040132344],"study_design_scores_gemma":[0.000022747363,0.000009910216,0.0009843172,0.00038541932,0.000010916344,0.000056122662,0.00010027948,0.00036542237,0.00062238565,0.003730955,0.99367696,0.00003454262],"about_ca_topic_score_codex":0.06307982,"about_ca_topic_score_gemma":0.13602854,"teacher_disagreement_score":0.98762363,"about_ca_system_score_codex":0.0040126634,"about_ca_system_score_gemma":0.014535627,"threshold_uncertainty_score":0.93816173},"labels":[],"label_agreement":null},{"id":"W3045889620","doi":"10.2196/17376","title":"Implementation of a Cohort Retrieval System for Clinical Data Repositories Using the Observational Medical Outcomes Partnership Common Data Model: Proof-of-Concept System Validation","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Center for Advancing Translational Sciences; National Institute of Biomedical Imaging and Bioengineering; National Institutes of Health","keywords":"Computer science; Information retrieval; Concept search; Unstructured data; Software portability; Scalability; Data mining; Database; Big data; Search engine; Web search query","score_opus":0.3134486461483542,"score_gpt":0.4776412865090831,"score_spread":0.1641926403607289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3045889620","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23383807,0.0004095088,0.66340816,0.0020137054,0.0003166395,0.00957525,0.005688642,0.08168148,0.0030685142],"genre_scores_gemma":[0.35032213,0.00022628071,0.6328604,0.0006387696,0.00005499214,0.004041655,0.009184917,0.00097313407,0.0016976817],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9898279,0.003578961,0.0012915371,0.0015543458,0.0033459342,0.00040128568],"domain_scores_gemma":[0.97371596,0.013593442,0.0013083243,0.004764801,0.005820592,0.00079679885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025744984,0.0009785127,0.0008467751,0.0014562634,0.0008822645,0.0024267547,0.0029840064,0.0014040418,0.0032985038],"category_scores_gemma":[0.052457575,0.0006063201,0.0011444298,0.0006838519,0.0009492944,0.0034290468,0.0032377306,0.0016989068,0.0015865021],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006349122,0.0066384054,0.064777166,0.0049438523,0.0012700369,0.003061901,0.005569892,0.08103213,0.17414048,0.020083554,0.058576494,0.5735569],"study_design_scores_gemma":[0.0036134068,0.0049455385,0.02676814,0.0004999205,0.00047051118,0.0013321339,0.0012338014,0.6597521,0.2421272,0.0062391935,0.05253953,0.0004786302],"about_ca_topic_score_codex":0.0070719537,"about_ca_topic_score_gemma":0.0044897418,"teacher_disagreement_score":0.025744984,"about_ca_system_score_codex":0.0017858504,"about_ca_system_score_gemma":0.0064195264,"threshold_uncertainty_score":0.13615417},"labels":[],"label_agreement":null},{"id":"W3048030338","doi":"10.20944/preprints202008.0220.v1","title":"The PHA4GE SARS-CoV-2 Contextual Data Specification for Open Genomic Epidemiology","year":2020,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University; Dalhousie University; McMaster University; BC Centre for Disease Control; University of British Columbia","funders":"Biotechnology and Biological Sciences Research Council","keywords":"Interoperability; Computer science; Openness to experience; Data science; Consistency (knowledge bases); Open science; Open data; Data integration; Reuse; Standardization; World Wide Web; Best practice; Alliance; Knowledge management; Data mining; Engineering; Geography; Political science","score_opus":0.5666488393984701,"score_gpt":0.48111947820519646,"score_spread":0.08552936119327365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3048030338","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004898271,0.00036771048,0.9217665,0.003075964,0.0006366942,0.0015139763,0.037047774,0.016803319,0.013889855],"genre_scores_gemma":[0.03605016,0.00086620514,0.8321033,0.0029726338,0.00033840854,0.0030802737,0.11090487,0.006519404,0.007164765],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9830451,0.005877743,0.00494398,0.0014265097,0.0037338377,0.000972866],"domain_scores_gemma":[0.9755491,0.0068115527,0.0019403206,0.008906097,0.005696457,0.0010964584],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.022136837,0.0012639903,0.0010018565,0.0033793128,0.0017082406,0.0060379105,0.002946572,0.003094735,0.006120637],"category_scores_gemma":[0.032776374,0.0014297311,0.0023665281,0.0031045426,0.0018718605,0.0047183065,0.006695372,0.004003496,0.007538178],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006125687,0.00026492588,0.009813009,0.0018460321,0.00015347896,0.000950837,0.0029197359,0.013190003,0.0132448105,0.61018586,0.2512406,0.09557818],"study_design_scores_gemma":[0.000059603008,0.000060547798,0.0015708986,0.00087307394,0.00006219747,0.00046816838,0.000514696,0.0103011,0.0074343365,0.070321396,0.908231,0.00010294381],"about_ca_topic_score_codex":0.013597191,"about_ca_topic_score_gemma":0.011412054,"teacher_disagreement_score":0.99705344,"about_ca_system_score_codex":0.0026002997,"about_ca_system_score_gemma":0.0110054305,"threshold_uncertainty_score":0.117072165},"labels":[],"label_agreement":null},{"id":"W305879088","doi":"","title":"Domain-specific synonym expansion and validation for biomedical information retrieval (multitext experiments for trec 2004)","year":2004,"lang":"en","type":"article","venue":"Text REtrieval Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":45,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Synonym (taxonomy); Information retrieval; Domain (mathematical analysis); Natural language processing; Artificial intelligence; Biology; Mathematics","score_opus":0.03568668182173386,"score_gpt":0.3012554734088946,"score_spread":0.26556879158716074,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W305879088","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.87501395,0.0036027723,0.063949175,0.0031722076,0.001272297,0.001874725,0.032941703,0.009107155,0.00906605],"genre_scores_gemma":[0.6908451,0.0009051308,0.19384436,0.0007658031,0.0002264772,0.0015258394,0.099846795,0.0013761182,0.010664391],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98125756,0.01281798,0.0015011033,0.0013980997,0.002475138,0.0005500527],"domain_scores_gemma":[0.95280796,0.030280583,0.0013469429,0.006546598,0.007842772,0.0011751325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020955298,0.0015426236,0.0012532808,0.0033305413,0.0028157977,0.0015981689,0.0017416253,0.0021405416,0.004920068],"category_scores_gemma":[0.0456827,0.0007402484,0.001496863,0.0028136882,0.0013576741,0.005033787,0.0027395138,0.0024524166,0.0033244702],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0097623635,0.0075160367,0.023827126,0.0029530644,0.0013877233,0.0011006376,0.003097933,0.02118467,0.14367022,0.0044240323,0.2849147,0.4961614],"study_design_scores_gemma":[0.005341374,0.0072809975,0.11539855,0.00035216022,0.002494431,0.0029110687,0.0036915636,0.3371896,0.42836654,0.00904293,0.08718477,0.00074595417],"about_ca_topic_score_codex":0.017707776,"about_ca_topic_score_gemma":0.018397914,"teacher_disagreement_score":0.020955298,"about_ca_system_score_codex":0.0016404197,"about_ca_system_score_gemma":0.0030354476,"threshold_uncertainty_score":0.11082351},"labels":[],"label_agreement":null},{"id":"W3082325268","doi":"10.1186/s12911-020-01227-6","title":"Identification of most influential co-occurring gene suites for gastrointestinal cancer using biomedical literature mining and graph-based influence maximization","year":2020,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Asia University; Ministry of Science and Technology, Taiwan","keywords":"Context (archaeology); Pipeline (software); Gastrointestinal cancer; Graph; Computer science; Computational biology; Cancer; Medicine; Colorectal cancer; Biology; Theoretical computer science; Internal medicine","score_opus":0.030386475622022117,"score_gpt":0.342508913807284,"score_spread":0.31212243818526186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3082325268","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40364236,0.009597235,0.5558598,0.0024855805,0.000121286284,0.0009831439,0.015587735,0.0036170878,0.008105718],"genre_scores_gemma":[0.6919598,0.0023135964,0.28949946,0.0003129309,0.0001444804,0.00042809595,0.0140646985,0.00020985738,0.0010670435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985127,0.00028039597,0.00012795355,0.00054632395,0.00043521327,0.00009749188],"domain_scores_gemma":[0.99516946,0.002896334,0.00072958175,0.00023668393,0.0007862698,0.0001816502],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001722418,0.0009741571,0.0008642178,0.01249347,0.0009608627,0.0013554939,0.0010090979,0.0008273697,0.001201034],"category_scores_gemma":[0.008818829,0.00031506692,0.0022217669,0.006172975,0.0006439921,0.0010438348,0.0009189064,0.0006252924,0.00051641103],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009395113,0.00072290906,0.25423345,0.0045894184,0.0023149676,0.0044972873,0.0015190892,0.112278536,0.070250556,0.015078348,0.019421158,0.5141548],"study_design_scores_gemma":[0.00011943813,0.00027399434,0.08943113,0.00045237102,0.0020152158,0.003024727,0.0006975244,0.795619,0.029731652,0.05389871,0.024591895,0.0001443905],"about_ca_topic_score_codex":0.0076706973,"about_ca_topic_score_gemma":0.013842569,"teacher_disagreement_score":0.01249347,"about_ca_system_score_codex":0.0013090582,"about_ca_system_score_gemma":0.002372895,"threshold_uncertainty_score":0.015252113},"labels":[],"label_agreement":null},{"id":"W3082742154","doi":"10.1158/1538-7445.am2020-3222","title":"Abstract 3222: The Virtual Molecular Tumor Board of the Variant Interpretation for Cancer Consortium: A systematic gateway connecting cancer genome interpretation and progress in genomic knowledgebases in cancer","year":2020,"lang":"en","type":"article","venue":"Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; University of Toronto","funders":"","keywords":"Harmonization; Context (archaeology); Precision medicine; Personalized medicine; Interpretation (philosophy); Medicine; Bioinformatics; Computer science; Biology; Pathology","score_opus":0.04652983227197934,"score_gpt":0.37830021205670594,"score_spread":0.3317703797847266,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3082742154","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03175548,0.0038230894,0.48713794,0.15974559,0.011220443,0.009595191,0.046762723,0.08075143,0.16920811],"genre_scores_gemma":[0.17086953,0.0017913079,0.59327585,0.021216238,0.006078774,0.007839146,0.101203725,0.030536031,0.06718944],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.93627405,0.044897236,0.0037041581,0.0031993336,0.010180211,0.001744992],"domain_scores_gemma":[0.8399306,0.059262507,0.007562655,0.03547775,0.030215012,0.027551541],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.1172414,0.0008792832,0.00087581645,0.006350477,0.003611875,0.012430399,0.0042133704,0.0039380323,0.04954534],"category_scores_gemma":[0.1375115,0.000954864,0.0007868449,0.00407993,0.0031063906,0.008907152,0.019100817,0.0037417302,0.02092422],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040417616,0.00017137447,0.003499954,0.0005792157,0.0000591674,0.00064543204,0.0043437784,0.0014293087,0.0038130896,0.017160606,0.8489011,0.11899272],"study_design_scores_gemma":[0.00022383929,0.00011741484,0.0028773388,0.0008299809,0.000036938538,0.00027057217,0.0019403801,0.004020866,0.0020056493,0.014170729,0.97340155,0.000104761486],"about_ca_topic_score_codex":0.0048986413,"about_ca_topic_score_gemma":0.008449321,"teacher_disagreement_score":0.1172414,"about_ca_system_score_codex":0.0032736622,"about_ca_system_score_gemma":0.021024982,"threshold_uncertainty_score":0.6200392},"labels":[],"label_agreement":null},{"id":"W3083382563","doi":"10.1002/ca.23677","title":"A tale of two systems: The tracts of my tears","year":2020,"lang":"en","type":"review","venue":"Clinical Anatomy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Organ system; Medicine; Alimentary tract; Meaning (existential); Nomenclature; Anatomy; Reproductive tract; Urinary system; Cognitive science; Physiology; Pathology; Epistemology; Biology; Zoology; Psychology; Philosophy; Taxonomy (biology); Internal medicine","score_opus":0.08970601884394448,"score_gpt":0.44917733069260274,"score_spread":0.35947131184865827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3083382563","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00086155935,0.86995256,0.00426981,0.07298802,0.009183483,0.00003310884,0.00016391199,0.00011978249,0.042427707],"genre_scores_gemma":[0.029877776,0.81355375,0.009160962,0.06363873,0.0072263065,0.00023108671,0.00042982656,0.00034320468,0.07553835],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99845266,0.000685721,0.00016858408,0.00020017546,0.0003702999,0.00012259152],"domain_scores_gemma":[0.9985482,0.00074329844,0.00014009452,0.00011373712,0.00027039793,0.00018421926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025755884,0.0008802627,0.001106242,0.0028408843,0.0025221426,0.006616139,0.0012525445,0.0030084448,0.015914693],"category_scores_gemma":[0.0054695136,0.00041293883,0.0008945244,0.0027676905,0.008014693,0.018853363,0.005475436,0.0072775413,0.0058229924],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013437346,0.00003689284,0.00035713142,0.006186794,0.00009033228,0.0009052417,0.008068946,0.00013117361,0.00076654326,0.27341494,0.36651963,0.34338808],"study_design_scores_gemma":[0.0000052787705,0.000012654819,0.00012725283,0.0023469846,0.000010581173,0.000979509,0.0010120346,0.000015553629,0.00006645243,0.0134325195,0.9819817,0.000009531439],"about_ca_topic_score_codex":0.0032342474,"about_ca_topic_score_gemma":0.0049017225,"teacher_disagreement_score":0.015914693,"about_ca_system_score_codex":0.00231183,"about_ca_system_score_gemma":0.005561738,"threshold_uncertainty_score":0.05324},"labels":[],"label_agreement":null},{"id":"W3084984950","doi":"10.1089/bio.2020.29075.drc","title":"Expanding the BBMRI-ERIC Directory into a Global Catalogue of COVID-19–Ready Collections: A Joint Initiative of BBMRI-ERIC and ISBER","year":2020,"lang":"en","type":"article","venue":"Biopreservation and Biobanking","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Directory; Coronavirus disease 2019 (COVID-19); Library science; 2019-20 coronavirus outbreak; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Joint (building); Geography; Political science; Biology; Medicine; Virology; Computer science; Engineering; Outbreak","score_opus":0.10486927818798714,"score_gpt":0.3234981186043995,"score_spread":0.21862884041641234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3084984950","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002967151,0.06330076,0.024903726,0.12310247,0.038424924,0.007650665,0.30450374,0.016975988,0.41817054],"genre_scores_gemma":[0.00474231,0.06420918,0.08952103,0.017774805,0.01601587,0.0068829823,0.39052054,0.014304232,0.396029],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.979506,0.004487737,0.0052314573,0.0015460205,0.007585149,0.0016436137],"domain_scores_gemma":[0.8412318,0.021775834,0.01642451,0.01467936,0.08080138,0.025087178],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.04460223,0.0017241726,0.0046618693,0.049895577,0.0034936122,0.021905815,0.0060323556,0.0037539797,0.24153864],"category_scores_gemma":[0.087284364,0.001744452,0.0017008245,0.053825453,0.0016446711,0.014618088,0.012286774,0.00393281,0.3023066],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000045979774,0.000012641267,0.00018657593,0.0008705675,0.000008986537,0.00004880087,0.00009167458,0.000032246946,0.00020101918,0.0010944962,0.95138586,0.046021204],"study_design_scores_gemma":[0.00000941107,0.00000630338,0.000496906,0.0006409747,0.000005131473,0.000043838456,0.000080527134,0.000013648412,0.00005064621,0.00018075253,0.9984571,0.0000146965685],"about_ca_topic_score_codex":0.01194359,"about_ca_topic_score_gemma":0.017292913,"teacher_disagreement_score":0.97809416,"about_ca_system_score_codex":0.005853073,"about_ca_system_score_gemma":0.03305373,"threshold_uncertainty_score":0.80802727},"labels":[],"label_agreement":null},{"id":"W3088981549","doi":"10.2196/18287","title":"Construction of a Digestive System Tumor Knowledge Graph Based on Chinese Electronic Medical Records: Development and Usability Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":28,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Schema (genetic algorithms); Readability; Knowledge extraction; Graph; Information retrieval; Usability; Knowledge graph; Natural language processing; Data mining; Artificial intelligence; Theoretical computer science","score_opus":0.011022309766078632,"score_gpt":0.28296049872643536,"score_spread":0.2719381889603567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3088981549","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.843468,0.0006580726,0.13210459,0.0011010214,0.00008440713,0.0044175694,0.0057579665,0.0053761923,0.007032352],"genre_scores_gemma":[0.6315493,0.0008180141,0.35421365,0.00024428213,0.00002018787,0.0014315387,0.009597061,0.0003257875,0.0018001023],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973786,0.001285215,0.00038268586,0.00041672963,0.0004595963,0.000077114586],"domain_scores_gemma":[0.9816699,0.0126586715,0.00075269054,0.0015761363,0.003012299,0.00033031107],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0071839257,0.00066021713,0.0005116996,0.005102538,0.0007946043,0.0015624483,0.0011764388,0.0005595616,0.0029140715],"category_scores_gemma":[0.019262299,0.00041652986,0.0011586447,0.0041507254,0.0005152571,0.003813715,0.0016340322,0.0006428779,0.00038905084],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00072752376,0.0019512771,0.07617698,0.0063838307,0.0004593472,0.0024304115,0.03298339,0.011132026,0.02079666,0.007854417,0.017004462,0.8220997],"study_design_scores_gemma":[0.0012364931,0.0046380544,0.2877362,0.0037779717,0.003361504,0.00599238,0.0753145,0.33785444,0.0605183,0.01892268,0.19982076,0.00082680763],"about_ca_topic_score_codex":0.009477782,"about_ca_topic_score_gemma":0.01135288,"teacher_disagreement_score":0.009477782,"about_ca_system_score_codex":0.0014506564,"about_ca_system_score_gemma":0.0023698842,"threshold_uncertainty_score":0.037992656},"labels":[],"label_agreement":null},{"id":"W3088992161","doi":"10.1016/j.annonc.2020.08.1696","title":"1382P Automating access to real-world evidence","year":2020,"lang":"en","type":"article","venue":"Annals of Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; Princess Margaret Cancer Centre","funders":"EMD Serono; Servier; Merck KGaA; PharmaMar; Bristol-Myers Squibb","keywords":"Concordance; Medicine; Data extraction; Lung cancer; Data collection; Stage (stratigraphy); Artificial intelligence; Medical physics; MEDLINE; Oncology; Internal medicine; Computer science; Statistics","score_opus":0.29968038010396375,"score_gpt":0.48459280946797706,"score_spread":0.1849124293640133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3088992161","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010170769,0.0012253531,0.6102095,0.0031703142,0.00039263236,0.0006520269,0.110503666,0.24558802,0.018087737],"genre_scores_gemma":[0.13620538,0.0014105083,0.668609,0.0012276233,0.000194695,0.0007682809,0.1668215,0.014268813,0.010494151],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981073,0.00037324167,0.00021326808,0.000431619,0.00077112153,0.0001034312],"domain_scores_gemma":[0.99207413,0.005126023,0.00027638965,0.0014736387,0.00089157483,0.0001582339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026523771,0.001892897,0.00090386096,0.0057533607,0.0006948383,0.0035524957,0.0015846639,0.001750521,0.036365308],"category_scores_gemma":[0.027607841,0.0008657257,0.0018863147,0.004631571,0.00058914424,0.0038572538,0.004419192,0.0012628976,0.013908572],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008507116,0.00017745583,0.0053145597,0.0030122946,0.0005727613,0.0015539583,0.0006175103,0.013352306,0.010972281,0.025311979,0.38376316,0.554501],"study_design_scores_gemma":[0.00034438085,0.00008451344,0.0052144546,0.000642929,0.00031843735,0.0013775275,0.00042782037,0.33515096,0.024723375,0.12558676,0.50601965,0.00010922333],"about_ca_topic_score_codex":0.010694434,"about_ca_topic_score_gemma":0.014689388,"teacher_disagreement_score":0.036365308,"about_ca_system_score_codex":0.0011683644,"about_ca_system_score_gemma":0.0024315605,"threshold_uncertainty_score":0.12165403},"labels":[],"label_agreement":null},{"id":"W3089809965","doi":"10.24095/hpcdp.29.3.02f","title":"Validité des diagnostics d’autisme recensés à l’aide de données administratives sur la santé","year":2009,"lang":"fr","type":"article","venue":"Maladies chroniques au Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"Dalhousie University","funders":"Canadian Institutes of Health Research; Dalhousie University; Autism Speaks","keywords":"Humanities; Political science; Physics; Philosophy","score_opus":0.0195331329185796,"score_gpt":0.27275354672369106,"score_spread":0.25322041380511146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3089809965","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29764098,0.010435421,0.5758225,0.02432636,0.0017193028,0.005316761,0.035917,0.009375033,0.039446697],"genre_scores_gemma":[0.41608,0.002639131,0.54970396,0.0028416861,0.00033876867,0.0029146092,0.017098567,0.00052395236,0.007859316],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.91989154,0.036887955,0.01227125,0.0104119815,0.018253159,0.002284094],"domain_scores_gemma":[0.70199955,0.19253972,0.023074666,0.022601854,0.057341646,0.0024425779],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.06808472,0.0016456791,0.0021538925,0.009396983,0.0025174231,0.008965884,0.0039417953,0.0021946975,0.005465821],"category_scores_gemma":[0.25129998,0.0013974459,0.0025349474,0.005855068,0.0017895875,0.006091768,0.0049584345,0.0033111759,0.0025496788],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016472547,0.00045713573,0.43641046,0.005285113,0.0017212655,0.00060046907,0.012011424,0.018318899,0.0061000506,0.022386108,0.03185116,0.46321055],"study_design_scores_gemma":[0.0008037325,0.0010473254,0.306411,0.013731752,0.0030326385,0.002921755,0.016203765,0.20544039,0.05102496,0.086106315,0.31225738,0.0010191156],"about_ca_topic_score_codex":0.039263256,"about_ca_topic_score_gemma":0.031729866,"teacher_disagreement_score":0.96073675,"about_ca_system_score_codex":0.0047774455,"about_ca_system_score_gemma":0.015791882,"threshold_uncertainty_score":0.3600707},"labels":[],"label_agreement":null},{"id":"W3094174332","doi":"10.1101/2020.10.17.20214460","title":"Applied Ontologies for Global Health Surveillance and Pandemic Intelligence","year":2020,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; University of New Brunswick","funders":"","keywords":"Interoperability; Data science; Pandemic; Computer science; Relevance (law); Interpretability; Ontology; Data integration; Preparedness; Knowledge management; Coronavirus disease 2019 (COVID-19); Political science; World Wide Web; Medicine; Data mining; Infectious disease (medical specialty); Artificial intelligence","score_opus":0.05572210080995837,"score_gpt":0.3537892742446846,"score_spread":0.2980671734347262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094174332","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011276299,0.004526886,0.9337462,0.022275185,0.00061441335,0.00036419206,0.002194925,0.0019979617,0.023003858],"genre_scores_gemma":[0.23069283,0.0050879964,0.7510295,0.0019062536,0.0005730952,0.00035912785,0.005291396,0.0003327877,0.0047271024],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9917562,0.004626153,0.0008550413,0.0006778757,0.001836164,0.00024865638],"domain_scores_gemma":[0.98925143,0.005830318,0.0008003606,0.0025024477,0.0012738624,0.0003415619],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013101886,0.0006246051,0.00066114403,0.0045474735,0.0016278682,0.007285652,0.0011768983,0.0020762605,0.004363635],"category_scores_gemma":[0.01912087,0.00042933592,0.0016189951,0.005342364,0.0023140945,0.00887403,0.0050353324,0.0021722575,0.00091139897],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038412607,0.00007380662,0.0015631391,0.00044953558,0.00014824873,0.0002780897,0.00076235825,0.015162418,0.0011252981,0.86593556,0.014394779,0.10006836],"study_design_scores_gemma":[0.000013545314,0.000009779164,0.00077027694,0.0003361649,0.00003717773,0.00012770008,0.00060333125,0.08290581,0.0013825333,0.81333697,0.100452445,0.000024301662],"about_ca_topic_score_codex":0.008100766,"about_ca_topic_score_gemma":0.00530355,"teacher_disagreement_score":0.013101886,"about_ca_system_score_codex":0.0036139179,"about_ca_system_score_gemma":0.0042070956,"threshold_uncertainty_score":0.06929022},"labels":[],"label_agreement":null},{"id":"W3094376906","doi":"10.2196/23449","title":"Searching PubMed to Retrieve Publications on the COVID-19 Pandemic: Comparative Analysis of Search Strings","year":2020,"lang":"en","type":"article","venue":"Journal of Medical Internet Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Generalitat de Catalunya; Centres de Recerca de Catalunya; Novo Nordisk Fonden; Ministerio de Ciencia, Innovación y Universidades","keywords":"Computer science; Information retrieval; Coronavirus disease 2019 (COVID-19); Sensitivity (control systems); Pandemic; String (physics); MEDLINE; Data mining; Medicine; Mathematics; Disease; Pathology; Biology","score_opus":0.3914811259219379,"score_gpt":0.5028015615962662,"score_spread":0.11132043567432831,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094376906","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6800426,0.2371024,0.012868689,0.004413249,0.0008834477,0.0074446844,0.04565349,0.0018504441,0.009741006],"genre_scores_gemma":[0.78681904,0.08954056,0.08011358,0.0020749052,0.00053218094,0.006944461,0.0318907,0.00063942216,0.0014451533],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.90943545,0.03723026,0.03666708,0.0051716706,0.010033441,0.0014621547],"domain_scores_gemma":[0.28931084,0.66239643,0.028140357,0.0065362337,0.012302331,0.0013138732],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06328839,0.0021053897,0.0048385626,0.049223896,0.0013372711,0.004882168,0.0020394712,0.0022806246,0.0045648003],"category_scores_gemma":[0.38678917,0.00083578646,0.0062271585,0.04215526,0.0013903215,0.008492423,0.0039222366,0.00097267085,0.001317838],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.018722745,0.00073294534,0.18634467,0.34953934,0.029916992,0.0036435025,0.010065211,0.00516233,0.011584826,0.0021567636,0.0182323,0.36389837],"study_design_scores_gemma":[0.005179719,0.013004015,0.5210795,0.15565926,0.10171608,0.019432265,0.020520363,0.030149186,0.020464292,0.014961472,0.095979415,0.0018545341],"about_ca_topic_score_codex":0.0020683943,"about_ca_topic_score_gemma":0.005503746,"teacher_disagreement_score":0.9367116,"about_ca_system_score_codex":0.0018685083,"about_ca_system_score_gemma":0.0052336557,"threshold_uncertainty_score":0.334705},"labels":[],"label_agreement":null},{"id":"W3095617614","doi":"10.2196/22898","title":"Extraction of Family History Information From Clinical Notes: Deep Learning and Heuristics Approach","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Relationship extraction; Information extraction; Heuristics; Computer science; Artificial intelligence; Task (project management); Natural language processing; Information retrieval; Machine learning; Data extraction; Data mining; MEDLINE; Engineering","score_opus":0.044762265218179484,"score_gpt":0.3279625408580742,"score_spread":0.2832002756398947,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3095617614","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18307562,0.0026693472,0.79961187,0.0013424247,0.00011153614,0.00036768845,0.0025877825,0.005310955,0.0049227704],"genre_scores_gemma":[0.670485,0.0009818439,0.31872493,0.0005313954,0.000086418375,0.00017005956,0.0048833787,0.00009265751,0.0040442836],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99954057,0.000103531136,0.000051173676,0.00016582913,0.0000751248,0.00006387463],"domain_scores_gemma":[0.9989887,0.000658795,0.000084575186,0.00007820807,0.00014210233,0.000047703885],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000843491,0.00086439034,0.00052730046,0.0026423843,0.0003500596,0.000878612,0.0010607282,0.0009272913,0.0014297202],"category_scores_gemma":[0.002528034,0.00037359772,0.0007776075,0.0015276602,0.00033029882,0.0010589971,0.0006950546,0.0008230445,0.00053938397],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033731584,0.0005001574,0.017456219,0.00027875518,0.00019926844,0.0007977453,0.00033516815,0.10368106,0.0072200857,0.0040020477,0.010026869,0.8551653],"study_design_scores_gemma":[0.000035360277,0.00008979755,0.004288265,0.000080265105,0.000116072086,0.00028926635,0.00013755591,0.9783102,0.005664776,0.007685192,0.0032787868,0.000024441615],"about_ca_topic_score_codex":0.013795319,"about_ca_topic_score_gemma":0.018787213,"teacher_disagreement_score":0.013795319,"about_ca_system_score_codex":0.0010091701,"about_ca_system_score_gemma":0.0013914773,"threshold_uncertainty_score":0.027430058},"labels":[],"label_agreement":null},{"id":"W3096045883","doi":"","title":"Functional Parthood: A Dispositional Perspective","year":2020,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Perspective (graphical); Computer science; Artificial intelligence","score_opus":0.019825158926965444,"score_gpt":0.24874822996874507,"score_spread":0.2289230710417796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096045883","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06149081,0.0040318435,0.6479467,0.020041559,0.00061060727,0.00008925211,0.0011102157,0.00035995163,0.26431906],"genre_scores_gemma":[0.91312796,0.0019387159,0.0630011,0.001489127,0.0007301501,0.00016935337,0.0010228365,0.00036134213,0.018159324],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9959716,0.0020454677,0.00024959567,0.0009102266,0.0004557596,0.0003673333],"domain_scores_gemma":[0.99280554,0.004079829,0.00041294977,0.0012100083,0.00088015455,0.0006115735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051431376,0.0006195478,0.00082399714,0.0024512724,0.0028293878,0.007386287,0.0020147448,0.0021103465,0.015868109],"category_scores_gemma":[0.008852174,0.00065601576,0.0013990765,0.0036408682,0.012009443,0.01788367,0.004203545,0.003291831,0.0016551937],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000014202677,0.0000057742864,0.000382521,0.000019220548,0.000007799327,0.00006457634,0.00080230366,0.00014641401,0.00012009635,0.9931465,0.00049723394,0.0047933855],"study_design_scores_gemma":[0.000006214854,0.000015358883,0.00051551015,0.000027840933,0.000013367121,0.0002513455,0.00069436687,0.00068538287,0.00014786821,0.98662066,0.011014244,0.000007722446],"about_ca_topic_score_codex":0.0019863974,"about_ca_topic_score_gemma":0.0011078591,"teacher_disagreement_score":0.015868109,"about_ca_system_score_codex":0.0017061115,"about_ca_system_score_gemma":0.0011299782,"threshold_uncertainty_score":0.053084075},"labels":[],"label_agreement":null},{"id":"W3096461368","doi":"10.2196/18953","title":"The Impact of Pretrained Language Models on Negation and Speculation Detection in Cross-Lingual Medical Text: Comparative Study","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":30,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Ministerio de Economía y Competitividad","keywords":"Computer science; Negation; Natural language processing; Artificial intelligence; Speculation; Conditional random field; Programming language","score_opus":0.03040752701534484,"score_gpt":0.3791577345735764,"score_spread":0.34875020755823155,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3096461368","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9324516,0.013573364,0.031700645,0.00198788,0.0011114985,0.0002546533,0.0045090583,0.004320554,0.010090833],"genre_scores_gemma":[0.9667009,0.0022246356,0.017121458,0.0003475093,0.0002068452,0.00015865512,0.010791914,0.00035379332,0.002094233],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99605983,0.0019513146,0.00035102954,0.0010291786,0.00034894512,0.00025971024],"domain_scores_gemma":[0.967477,0.027328921,0.00069058005,0.0018613556,0.0022264682,0.00041565328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008432997,0.0023850119,0.001263587,0.0019123427,0.00082406664,0.0029818593,0.001995193,0.0018325215,0.0037681623],"category_scores_gemma":[0.02796379,0.0007134499,0.0018109299,0.0013914626,0.00086933805,0.005645773,0.0021564725,0.0027945917,0.002896902],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063109994,0.0033652934,0.051857688,0.0031520273,0.0038157308,0.0020443797,0.001853326,0.13714151,0.00962503,0.0021955804,0.028134726,0.7505038],"study_design_scores_gemma":[0.00022140576,0.0010879012,0.021017257,0.00028841916,0.0012410212,0.0006634808,0.0013628601,0.9583394,0.007577683,0.003583259,0.00442481,0.00019251042],"about_ca_topic_score_codex":0.020203535,"about_ca_topic_score_gemma":0.016820459,"teacher_disagreement_score":0.020203535,"about_ca_system_score_codex":0.001786432,"about_ca_system_score_gemma":0.0013759172,"threshold_uncertainty_score":0.04459852},"labels":[],"label_agreement":null},{"id":"W3097585533","doi":"10.2196/22333","title":"Automatic Structuring of Ontology Terms Based on Lexical Granularity and Machine Learning: Algorithm Development and Validation","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Natural language processing; Task (project management); Artificial intelligence; Convolutional neural network; Ontology; Relation (database); Pairwise comparison; Word embedding; Granularity; Set (abstract data type); Structuring; Machine learning; Embedding; Data mining; Programming language","score_opus":0.020466560884490614,"score_gpt":0.27326704383340467,"score_spread":0.25280048294891405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097585533","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.048864074,0.00052101276,0.94268185,0.0002877097,0.000045199227,0.00032995787,0.000514366,0.0056365756,0.001119215],"genre_scores_gemma":[0.17298086,0.00019169279,0.82401645,0.000099620906,0.000026879836,0.00028000431,0.0014695053,0.00020813552,0.0007268074],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99879324,0.00036038063,0.00014296683,0.00036312436,0.00024820172,0.00009191533],"domain_scores_gemma":[0.9954709,0.0027746945,0.0003499428,0.00058064103,0.0007243699,0.00009947598],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00274685,0.0007828544,0.0007614572,0.0027519842,0.00058636366,0.0015113436,0.0014544618,0.0013263437,0.0029883632],"category_scores_gemma":[0.008581412,0.00034606436,0.0007875684,0.0020796522,0.0005933676,0.0021241999,0.00150537,0.001296639,0.0013325448],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021078726,0.00024016206,0.0054773185,0.00030529147,0.00009470639,0.000092598224,0.00020363014,0.058522332,0.014648814,0.0042128633,0.0050717266,0.91091985],"study_design_scores_gemma":[0.000038762173,0.000044009837,0.0015601692,0.00004379873,0.000025534591,0.00009922406,0.000112635855,0.97924167,0.009425315,0.00734066,0.002054751,0.00001354542],"about_ca_topic_score_codex":0.006299876,"about_ca_topic_score_gemma":0.007800872,"teacher_disagreement_score":0.006299876,"about_ca_system_score_codex":0.0013806053,"about_ca_system_score_gemma":0.0019898403,"threshold_uncertainty_score":0.014526904},"labels":[],"label_agreement":null},{"id":"W3097926973","doi":"10.1093/ajcp/aqaa137.033","title":"What’s in a Name? Comparative Analysis of Laboratory Test Naming Guidelines as Applied to Common Confusing Test Names","year":2020,"lang":"en","type":"article","venue":"American Journal of Clinical Pathology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Computer science; Respondent; Psychology","score_opus":0.08670043127588939,"score_gpt":0.44350665475596635,"score_spread":0.35680622348007696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3097926973","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9887028,0.00059262785,0.00460492,0.0016247133,0.00009161426,0.0003891713,0.0002187854,0.00005868034,0.00371662],"genre_scores_gemma":[0.993206,0.0005010709,0.004581666,0.0006328002,0.000025554651,0.00047775032,0.00021870511,0.000042505366,0.00031393318],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9140311,0.053707477,0.012100498,0.002871664,0.014522211,0.0027671596],"domain_scores_gemma":[0.5358311,0.31063825,0.07513641,0.00896526,0.06539632,0.004032663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.066263266,0.00030390176,0.00070129364,0.004898474,0.0016754502,0.0045812754,0.0016410601,0.0015791412,0.0020593766],"category_scores_gemma":[0.39314607,0.0004033127,0.00087973115,0.0044570607,0.003533817,0.0075954855,0.004589725,0.0017779869,0.0004920611],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008120383,0.00039316676,0.5026947,0.0022585778,0.0001502645,0.0009972646,0.36440438,0.00073020265,0.0025094051,0.0063031684,0.006228837,0.11251803],"study_design_scores_gemma":[0.00006621064,0.00090821856,0.4099631,0.0029183752,0.00021130894,0.0010406343,0.55126214,0.0059360717,0.002756314,0.002925559,0.021733154,0.00027895428],"about_ca_topic_score_codex":0.009367074,"about_ca_topic_score_gemma":0.010565264,"teacher_disagreement_score":0.066263266,"about_ca_system_score_codex":0.0056864894,"about_ca_system_score_gemma":0.00909287,"threshold_uncertainty_score":0.35043782},"labels":[],"label_agreement":null},{"id":"W3098769586","doi":"10.1371/journal.pone.0242353","title":"The recipes of Philosophy of Science: Characterizing the semantic structure of corpora by means of topic associative rules","year":2020,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"Fonds de Recherche du Québec-Société et Culture; Canada Research Chairs; Canada Foundation for Innovation; Université du Québec à Montréal","keywords":"Computer science; Associative property; Identification (biology); Information retrieval; Set (abstract data type); Natural language processing; Ontology; Topic model; Semantic similarity; Artificial intelligence; Data science; Epistemology","score_opus":0.036291322330425244,"score_gpt":0.24137936873862462,"score_spread":0.2050880464081994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3098769586","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28530183,0.0020475881,0.698458,0.0011014929,0.00007475265,0.00046851006,0.0039076577,0.0009331979,0.0077069234],"genre_scores_gemma":[0.6225855,0.0010430495,0.36646032,0.00017714592,0.00014682538,0.0012158045,0.007163905,0.00026356845,0.0009439348],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99341995,0.003411811,0.0006000643,0.0015662684,0.0008453342,0.00015659166],"domain_scores_gemma":[0.9583027,0.031144354,0.0035548443,0.004449242,0.0020958346,0.0004530205],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.008375213,0.00069979817,0.0007400686,0.008640175,0.0014937188,0.004664474,0.0012029049,0.0011217127,0.0014776159],"category_scores_gemma":[0.05083924,0.00077323196,0.0013367286,0.008703453,0.0026310831,0.0076938337,0.0019490903,0.0017113367,0.00057175395],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005428673,0.00049620494,0.18121417,0.0019496897,0.0009320509,0.0010515129,0.021925012,0.05230575,0.01698093,0.21399873,0.009255644,0.49934748],"study_design_scores_gemma":[0.00009613405,0.00016810826,0.0905826,0.00042288276,0.0003422495,0.0009488661,0.0039762203,0.45974427,0.008914793,0.4057284,0.028905723,0.00016984025],"about_ca_topic_score_codex":0.0038755252,"about_ca_topic_score_gemma":0.0052134814,"teacher_disagreement_score":0.9916248,"about_ca_system_score_codex":0.0011847115,"about_ca_system_score_gemma":0.0014032191,"threshold_uncertainty_score":0.044292867},"labels":[],"label_agreement":null},{"id":"W3099059156","doi":"10.1101/414136","title":"Increasing metadata coverage of SRA BioSample entries using deep learning based Named Entity Recognition","year":2018,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; Canadian Institute for Advanced Research","keywords":"Metadata; Computer science; Scalability; Classifier (UML); Information retrieval; Artificial neural network; Named-entity recognition; Artificial intelligence; Annotation; World Wide Web; Database","score_opus":0.02594117714831545,"score_gpt":0.2487920017763017,"score_spread":0.22285082462798625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3099059156","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8640954,0.0010895029,0.08513535,0.0008254865,0.00024405547,0.00010677236,0.025638225,0.019007662,0.0038575267],"genre_scores_gemma":[0.7873525,0.00031528683,0.13901845,0.00047513272,0.000085440704,0.00019361544,0.069304444,0.00048214576,0.0027730283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992442,0.00012584831,0.00007173814,0.00033555928,0.00015311275,0.00006961944],"domain_scores_gemma":[0.99739033,0.0010154303,0.00034409252,0.00042654105,0.0007038008,0.00011987209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016990403,0.0006816049,0.0004999247,0.0019476996,0.00048830005,0.0007577719,0.0009239528,0.000662204,0.0011835553],"category_scores_gemma":[0.0044880374,0.00019503933,0.00074628036,0.0010964846,0.00031084512,0.001511081,0.0011919184,0.00085814635,0.0015085592],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015875325,0.0008083485,0.20744711,0.0008076757,0.00043959141,0.0009923147,0.0008828599,0.073714964,0.18336833,0.0017533781,0.0405534,0.48764443],"study_design_scores_gemma":[0.000046250985,0.00030715056,0.052743755,0.0001186486,0.00015992568,0.000244671,0.00050853676,0.82399046,0.10388424,0.0030289635,0.014868293,0.00009906119],"about_ca_topic_score_codex":0.0058681974,"about_ca_topic_score_gemma":0.0099237,"teacher_disagreement_score":0.0058681974,"about_ca_system_score_codex":0.00046986854,"about_ca_system_score_gemma":0.00044585823,"threshold_uncertainty_score":0.011668086},"labels":[],"label_agreement":null},{"id":"W3101981777","doi":"10.1093/database/baaa079","title":"Measurement Recorder: developing a useful tool for making species descriptions that produces computable phenotypes","year":2020,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Agriculture and Agri-Food Canada; University of Manitoba","funders":"National Science Foundation","keywords":"Computer science; Usability; Software; Set (abstract data type); Ontology; Reuse; Variation (astronomy); Character (mathematics); Information retrieval; Data science; Software engineering; Human–computer interaction; Programming language","score_opus":0.17239052536197122,"score_gpt":0.3014869736680668,"score_spread":0.1290964483060956,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3101981777","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002117612,0.00008867311,0.8229253,0.00041662832,0.00012596705,0.00047757768,0.005045022,0.16623864,0.0025644852],"genre_scores_gemma":[0.018327292,0.00027269308,0.93472695,0.00032394473,0.00005176407,0.0013364516,0.010204919,0.029230118,0.0055258702],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99465066,0.0014386947,0.0009184904,0.0010107182,0.0018230913,0.00015846611],"domain_scores_gemma":[0.96754074,0.018385846,0.0023158563,0.0074496884,0.00347752,0.0008303328],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01597781,0.0020912865,0.0012905526,0.004860613,0.0012280213,0.0040742396,0.004606051,0.0018098991,0.026852971],"category_scores_gemma":[0.0458105,0.0026993272,0.0026672436,0.0037252763,0.0014680044,0.010542324,0.005288737,0.003289466,0.011710982],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001017628,0.00047450114,0.007987884,0.0045930594,0.0005447965,0.0014757321,0.008319071,0.00910005,0.04303045,0.06252945,0.22644223,0.6344852],"study_design_scores_gemma":[0.00042430515,0.00046578835,0.0051404396,0.00097265723,0.00020506972,0.0018281466,0.0011478007,0.08461258,0.07358461,0.056796275,0.7741722,0.000650087],"about_ca_topic_score_codex":0.0022901522,"about_ca_topic_score_gemma":0.0024157693,"teacher_disagreement_score":0.026852971,"about_ca_system_score_codex":0.0016500123,"about_ca_system_score_gemma":0.003352416,"threshold_uncertainty_score":0.08983213},"labels":[],"label_agreement":null},{"id":"W3102749286","doi":"10.2196/23104","title":"Clinical Term Normalization Using Learned Edit Patterns and Subconcept Matching: System Development and Evaluation","year":2020,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Normalization (sociology); Natural language processing; Edit distance; Unified Medical Language System; Artificial intelligence; Term (time); Information retrieval","score_opus":0.07249807901364434,"score_gpt":0.37381956233859637,"score_spread":0.301321483324952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3102749286","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39120054,0.0073580816,0.44660217,0.0021509198,0.0015445729,0.005871616,0.012185332,0.12530673,0.007780011],"genre_scores_gemma":[0.34694386,0.0016266794,0.60683304,0.0010755407,0.00021003881,0.0022739966,0.033947803,0.0018267179,0.0052623856],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99289227,0.0018455704,0.0009159276,0.002289452,0.0017890458,0.0002678363],"domain_scores_gemma":[0.98956853,0.0048081055,0.00052655133,0.0013826118,0.0032296337,0.0004844642],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007879043,0.0015798114,0.0016086724,0.0024406721,0.00092739554,0.0021163528,0.0037500262,0.0024097094,0.004957115],"category_scores_gemma":[0.025936786,0.0005349118,0.0010344663,0.0020697135,0.00070152595,0.0031889842,0.0024467444,0.0016542329,0.0034369081],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018512101,0.00161133,0.009256686,0.002342708,0.00063081394,0.0010090516,0.00082457176,0.0136250425,0.041459627,0.0010172257,0.048049927,0.8783217],"study_design_scores_gemma":[0.0014318095,0.002776477,0.021277292,0.00039318218,0.0007626401,0.004417958,0.00094953063,0.76005036,0.15671976,0.0045983023,0.046247266,0.0003755311],"about_ca_topic_score_codex":0.009306178,"about_ca_topic_score_gemma":0.007401658,"teacher_disagreement_score":0.009306178,"about_ca_system_score_codex":0.0018068573,"about_ca_system_score_gemma":0.003575843,"threshold_uncertainty_score":0.041668892},"labels":[],"label_agreement":null},{"id":"W3108384885","doi":"10.21203/rs.3.rs-40780/v1","title":"Decoding semi-automated title-abstract screening: a retrospective exploration of the review, study, and publication characteristics associated with accurate relevance predictions","year":2020,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Government of Canada; Agency for Healthcare Research and Quality; U.S. Department of Health and Human Services","keywords":"Relevance (law); Decoding methods; Computer science; Information retrieval; Data science; Data mining; Algorithm; Political science","score_opus":0.10737056613387296,"score_gpt":0.39861958659113605,"score_spread":0.2912490204572631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3108384885","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6126906,0.0926225,0.2094615,0.006998067,0.0010244658,0.009567576,0.058985557,0.0031477609,0.00550191],"genre_scores_gemma":[0.9096271,0.0047427756,0.06911783,0.001135699,0.00029835277,0.0050005107,0.009157445,0.00045863594,0.00046174097],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.64963806,0.20996346,0.092881545,0.017634979,0.028173294,0.0017087538],"domain_scores_gemma":[0.10808696,0.73397386,0.100903146,0.033424478,0.022783805,0.00082768913],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.26723015,0.001512229,0.0028777502,0.008844007,0.00072947727,0.003983405,0.001953626,0.0013210714,0.0029982831],"category_scores_gemma":[0.6949728,0.0014836737,0.0065231635,0.00873395,0.0012839718,0.0041141724,0.0026661898,0.0014852972,0.0012732933],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008723616,0.00018435296,0.60927474,0.05650199,0.017698469,0.0013866854,0.0033453046,0.010833227,0.003604951,0.0018893664,0.01749425,0.269063],"study_design_scores_gemma":[0.0037834567,0.005490696,0.657726,0.043341264,0.06762097,0.011793057,0.0019888924,0.096435346,0.021291388,0.017529236,0.07181421,0.0011855104],"about_ca_topic_score_codex":0.0018241918,"about_ca_topic_score_gemma":0.0032676603,"teacher_disagreement_score":0.73276985,"about_ca_system_score_codex":0.0016993799,"about_ca_system_score_gemma":0.0051960587,"threshold_uncertainty_score":0.9036357},"labels":[],"label_agreement":null},{"id":"W3109203371","doi":"10.5858/arpa.2020-0239-ed","title":"Exploring the College of American Pathologists Electronic Cancer Checklists: What They Are and What They Can Do for You","year":2020,"lang":"en","type":"editorial","venue":"Archives of Pathology & Laboratory Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trillium Health Centre","funders":"","keywords":"Vendor; Computer science; Software portability; Interoperability; Electronic data; Medicine; Medical physics; Data science; Information retrieval; World Wide Web","score_opus":0.021596570227801013,"score_gpt":0.291315873521626,"score_spread":0.269719303293825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109203371","genre_codex":"commentary","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09985477,0.02502649,0.06595888,0.70435196,0.005916966,0.0028021508,0.0022358086,0.0015521224,0.092300944],"genre_scores_gemma":[0.53846174,0.029340548,0.2680435,0.13727552,0.0030939437,0.0064078025,0.0027267775,0.0011071377,0.013543096],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.86668307,0.082395636,0.011582886,0.004407997,0.030274907,0.0046555507],"domain_scores_gemma":[0.66803205,0.18678059,0.030873148,0.014760987,0.08346869,0.016084619],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10892046,0.00077529973,0.0005981996,0.0068174424,0.005867733,0.015679192,0.0033569867,0.0044488716,0.0075074183],"category_scores_gemma":[0.3185186,0.0012925386,0.0010456099,0.0060090586,0.0081389705,0.0262611,0.0141282715,0.006951099,0.0021384037],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015123966,0.00024585496,0.044096146,0.0027146155,0.00006546027,0.0004253437,0.11953142,0.00036037344,0.00060731766,0.0455555,0.28478277,0.50146407],"study_design_scores_gemma":[0.000072502036,0.0002679869,0.01770541,0.009913139,0.000051446175,0.0010038816,0.124730006,0.00083367614,0.0005791251,0.03914057,0.80544895,0.00025329745],"about_ca_topic_score_codex":0.009147049,"about_ca_topic_score_gemma":0.016582297,"teacher_disagreement_score":0.10892046,"about_ca_system_score_codex":0.008939678,"about_ca_system_score_gemma":0.024564344,"threshold_uncertainty_score":0.57603335},"labels":[],"label_agreement":null},{"id":"W3110797946","doi":"","title":"QUESTO – An Ontology for Questionnaires","year":2020,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Ontology; Computer science; Interoperability; Context (archaeology); Information retrieval; Component (thermodynamics); Knowledge management; Data science; World Wide Web","score_opus":0.024312192267514677,"score_gpt":0.27726328939408873,"score_spread":0.25295109712657404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3110797946","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01059004,0.00044791857,0.8890447,0.0026368427,0.000357246,0.0031861342,0.05382739,0.023155175,0.016754478],"genre_scores_gemma":[0.09765472,0.0011099108,0.75919026,0.0018764,0.00021194243,0.0068082158,0.11260221,0.0043532476,0.016193168],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98973936,0.0035235942,0.0024660132,0.0013304083,0.0025279229,0.0004126679],"domain_scores_gemma":[0.9845926,0.008473099,0.0011158921,0.0028576266,0.0021286788,0.0008321772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010867515,0.0010803111,0.0013996136,0.0051117693,0.0015852631,0.004682447,0.0018023889,0.0016621343,0.01114258],"category_scores_gemma":[0.032558113,0.0010829932,0.0026318228,0.0049416134,0.0013502268,0.010107152,0.0055706063,0.0018234834,0.0068206624],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00077122304,0.0004937391,0.012491424,0.00337761,0.0002948321,0.00060678815,0.0055951807,0.0060336567,0.009687079,0.415183,0.12334876,0.42211676],"study_design_scores_gemma":[0.00011084161,0.000111473,0.0071670837,0.0008776789,0.00012103753,0.0005570568,0.001968446,0.029835735,0.004408193,0.22322854,0.73144287,0.00017113608],"about_ca_topic_score_codex":0.0069961753,"about_ca_topic_score_gemma":0.0050932537,"teacher_disagreement_score":0.01114258,"about_ca_system_score_codex":0.0025070214,"about_ca_system_score_gemma":0.0051140785,"threshold_uncertainty_score":0.05747366},"labels":[],"label_agreement":null},{"id":"W3111376437","doi":"10.5334/dsj-2020-047","title":"39 Hints to Facilitate the Use of Semantics for Data on Agriculture and Nutrition","year":2020,"lang":"en","type":"article","venue":"Data Science Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Agriculture and Agri-Food Canada; Rural Development Administration; Agence Nationale de la Recherche; Ministry of Agriculture of the People's Republic of China; Department for International Development; Bill and Melinda Gates Foundation","keywords":"Computer science; Semantic interoperability; Interoperability; Conceptualization; Linked data; Knowledge management; Data science; Semantics (computer science); World Wide Web; Data sharing; Standardization; Semantic Web","score_opus":0.29656193485514376,"score_gpt":0.34992526355676146,"score_spread":0.053363328701617696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111376437","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02093625,0.0012818517,0.90299493,0.030079639,0.00034743815,0.0017061528,0.005110662,0.0055394187,0.032003712],"genre_scores_gemma":[0.05460829,0.000724679,0.93190116,0.0024699501,0.000074507225,0.0006977443,0.0060099866,0.0009165869,0.0025971239],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9369188,0.03146348,0.014995774,0.0032604032,0.011978676,0.0013828344],"domain_scores_gemma":[0.8023643,0.121781625,0.0071572065,0.04042386,0.02545759,0.0028153826],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.084525734,0.001227408,0.0011606202,0.008987781,0.0025436287,0.010833569,0.0032143309,0.004687839,0.0049406476],"category_scores_gemma":[0.12998454,0.0017901405,0.0027164696,0.00771523,0.0058262013,0.024211813,0.014625782,0.0053029456,0.0029748136],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027551217,0.0001721869,0.008395811,0.002604551,0.00010542783,0.001401869,0.023798136,0.0045676497,0.016004065,0.7778568,0.025721394,0.13909657],"study_design_scores_gemma":[0.00003892743,0.000090799724,0.0024094526,0.0021670593,0.000075187,0.0009278917,0.0051509123,0.007037683,0.0091742,0.20733798,0.7654444,0.00014543238],"about_ca_topic_score_codex":0.008073497,"about_ca_topic_score_gemma":0.011843875,"teacher_disagreement_score":0.084525734,"about_ca_system_score_codex":0.0041249283,"about_ca_system_score_gemma":0.009599124,"threshold_uncertainty_score":0.44702017},"labels":[],"label_agreement":null},{"id":"W3111411783","doi":"10.23889/ijpds.v5i5.1637","title":"Deep Learning and NLP For Knowledge Extraction from Laboratory Reports","year":2020,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Vector Institute; University of Toronto","funders":"","keywords":"Computer science; Artificial intelligence; Identifier; Natural language processing; Named-entity recognition; Parsing; Information extraction; Information retrieval; F1 score; Deep learning; Identification (biology); ENCODE; Task (project management); Machine learning","score_opus":0.056942037700102754,"score_gpt":0.3996541144822478,"score_spread":0.34271207678214505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111411783","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020661736,0.0013707387,0.9502268,0.0016691363,0.00011179345,0.00039124247,0.00874243,0.014005015,0.0028211304],"genre_scores_gemma":[0.19159786,0.0010869666,0.7820077,0.00055130327,0.000110016335,0.00076022843,0.020021025,0.0002672557,0.0035976346],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997335,0.00086989877,0.0003673522,0.0007151867,0.0005186656,0.00019375092],"domain_scores_gemma":[0.9918412,0.006186613,0.0004990875,0.0005354006,0.00084709283,0.000090694724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035880662,0.0013128782,0.0007460364,0.004344754,0.0008661774,0.002011879,0.0018658601,0.0015763036,0.004804558],"category_scores_gemma":[0.012188626,0.00068025157,0.0015682926,0.0036136126,0.00079876615,0.0020857677,0.0018828587,0.0022015802,0.002044917],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015133878,0.0002250297,0.0039504925,0.0008469044,0.00018588256,0.0004620719,0.00027574357,0.1669276,0.005617655,0.007930282,0.018576603,0.7948504],"study_design_scores_gemma":[0.000022255299,0.000044850865,0.0011092542,0.00012020436,0.000041885574,0.0001029752,0.00009649328,0.95762104,0.008076314,0.021734627,0.011006057,0.000023987073],"about_ca_topic_score_codex":0.02096538,"about_ca_topic_score_gemma":0.019788265,"teacher_disagreement_score":0.02096538,"about_ca_system_score_codex":0.0023517713,"about_ca_system_score_gemma":0.0032340088,"threshold_uncertainty_score":0.041686654},"labels":[],"label_agreement":null},{"id":"W3111616091","doi":"10.23889/ijpds.v5i5.1540","title":"Using Reproducible Data Visualizations to Augment Decision-Making During Suppression of Small Counts","year":2020,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Workflow; Redaction; Computer science; Human error; Judgement; Data quality; Population; Data science; Medicine; Risk analysis (engineering); Database; Engineering; Operations management; Geography; Political science","score_opus":0.2227099826820303,"score_gpt":0.4798295732079889,"score_spread":0.2571195905259586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111616091","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.041638464,0.000582972,0.8412775,0.0051061534,0.00069035543,0.0014852146,0.008351708,0.094061814,0.0068059135],"genre_scores_gemma":[0.14789046,0.0002742581,0.8390471,0.0007274918,0.00019319908,0.0010462437,0.0052766996,0.003919439,0.0016251968],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9824691,0.008063654,0.0021206995,0.0033955905,0.0034255758,0.00052535237],"domain_scores_gemma":[0.86602926,0.07138372,0.011541022,0.031780593,0.016054366,0.0032109737],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.026848854,0.0024079138,0.001168197,0.004097541,0.0016764379,0.008869602,0.003974193,0.0016886974,0.0119479075],"category_scores_gemma":[0.09620449,0.0011233151,0.0021947725,0.0021239694,0.0020797926,0.0045300378,0.0067837248,0.002537116,0.0035759348],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031837416,0.00084169453,0.050173942,0.0041078837,0.00068993145,0.0018329906,0.017693052,0.044215042,0.045007456,0.030271862,0.09849386,0.70348865],"study_design_scores_gemma":[0.0008067215,0.001209952,0.03295965,0.0026502085,0.00043108687,0.001034985,0.003644617,0.4024859,0.11755471,0.13761124,0.29829764,0.0013133241],"about_ca_topic_score_codex":0.004338998,"about_ca_topic_score_gemma":0.005781247,"teacher_disagreement_score":0.97315115,"about_ca_system_score_codex":0.0018875737,"about_ca_system_score_gemma":0.0044496153,"threshold_uncertainty_score":0.14199197},"labels":[],"label_agreement":null},{"id":"W3111762657","doi":"10.23889/ijpds.v5i5.1621","title":"MASK: A Success Story for An International Collaboration","year":2020,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Identification (biology); Masking (illustration); Process (computing); Protected health information; Software; Interface (matter); Artificial intelligence; Machine learning; Information retrieval; Data science; World Wide Web; Public health","score_opus":0.09876678348695475,"score_gpt":0.4311233464971597,"score_spread":0.332356563010205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3111762657","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016600411,0.009218624,0.0024510354,0.8142506,0.01651179,0.00006973918,0.00073688413,0.0006835793,0.13947746],"genre_scores_gemma":[0.3758146,0.0094582485,0.0065178573,0.28034455,0.010383385,0.00036205084,0.001981964,0.0019400866,0.3131972],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9873707,0.005307384,0.00046265413,0.0012513857,0.0029022384,0.0027055244],"domain_scores_gemma":[0.9817892,0.0030007036,0.0008861831,0.001404196,0.0027804996,0.01013919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0148359,0.0009590344,0.00063334446,0.0012701004,0.013657413,0.016948303,0.0018688802,0.008793519,0.038287755],"category_scores_gemma":[0.021536624,0.000417172,0.00067042944,0.0017891412,0.0066994703,0.01773576,0.018225621,0.011637002,0.014571859],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013718312,0.00007531215,0.0018378774,0.00017364889,0.000023878922,0.0016652298,0.012479218,0.00015126569,0.00039316464,0.13397577,0.80096406,0.048123375],"study_design_scores_gemma":[0.000013464478,0.000035198653,0.0006788779,0.00017725675,0.00000461198,0.0004804256,0.0068564108,0.000117319556,0.00012036704,0.0075869886,0.98390377,0.00002527941],"about_ca_topic_score_codex":0.005372543,"about_ca_topic_score_gemma":0.005851439,"teacher_disagreement_score":0.038287755,"about_ca_system_score_codex":0.0045184623,"about_ca_system_score_gemma":0.008313596,"threshold_uncertainty_score":0.12808532},"labels":[],"label_agreement":null},{"id":"W3112885733","doi":"10.3389/fpubh.2020.515347","title":"Automatic Identification of Information Quality Metrics in Health News Stories","year":2020,"lang":"en","type":"article","venue":"Frontiers in Public Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"University of Brighton; College of Health, Education, and Human Development, Clemson University","keywords":"Computer science; Quality (philosophy); Artificial intelligence; Health care; Identification (biology); Publication; Process (computing); Machine learning; Task (project management); Natural language; Set (abstract data type); Natural language processing; Data science","score_opus":0.05770873701829786,"score_gpt":0.3433739866807531,"score_spread":0.2856652496624552,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3112885733","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.82345605,0.0053265174,0.13896383,0.002494034,0.0004657869,0.0011724022,0.014749987,0.0055493475,0.007821988],"genre_scores_gemma":[0.85024136,0.000523526,0.1289891,0.00012931932,0.00029547443,0.00043772653,0.018149886,0.00015893354,0.0010746482],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.991046,0.0032117127,0.0014287339,0.0014307642,0.0024857728,0.00039699263],"domain_scores_gemma":[0.9064232,0.062091433,0.012560278,0.002146324,0.015769579,0.0010090874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009143898,0.0010616741,0.000923386,0.0136287175,0.00063319696,0.0037114557,0.0009771028,0.0014202921,0.0011876521],"category_scores_gemma":[0.059851903,0.00034865894,0.0008251844,0.004104195,0.0006778848,0.003873525,0.0013381656,0.0013406721,0.00071699003],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031531975,0.0010221831,0.15987483,0.0049119005,0.00065711665,0.0012471393,0.0038787764,0.0145382425,0.041781638,0.003722521,0.028549612,0.736663],"study_design_scores_gemma":[0.00022656588,0.0011027916,0.29703468,0.0007423739,0.0005260299,0.0013675039,0.0040266705,0.603225,0.06407798,0.008153184,0.01926989,0.0002473074],"about_ca_topic_score_codex":0.0026226593,"about_ca_topic_score_gemma":0.0026788,"teacher_disagreement_score":0.0136287175,"about_ca_system_score_codex":0.0016392167,"about_ca_system_score_gemma":0.0009055897,"threshold_uncertainty_score":0.048358142},"labels":[],"label_agreement":null},{"id":"W3113186412","doi":"10.1186/s12911-020-01288-7","title":"Web-based interactive mapping from data dictionaries to ontologies, with an application to cancer registry","year":2020,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"National Cancer Institute; National Institutes of Health","keywords":"Computer science; Ontology; Interface (matter); Data dictionary; Information retrieval; Data mapping; User interface; Health informatics; Cancer registry; Data mining; Database; Cancer; Metadata; World Wide Web; Health care; Medicine","score_opus":0.05915755765944522,"score_gpt":0.35411830925454324,"score_spread":0.29496075159509805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3113186412","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.030398823,0.00017052409,0.7033926,0.0007618934,0.000086636326,0.0014311221,0.0032579019,0.2524183,0.008082174],"genre_scores_gemma":[0.10377681,0.00027937113,0.86288446,0.0005365494,0.000058019,0.001664374,0.008521193,0.012369021,0.00991017],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99853706,0.00042709173,0.00018153388,0.00034452873,0.00043072194,0.00007907866],"domain_scores_gemma":[0.9935226,0.0038307137,0.00026670165,0.0012249793,0.00079647545,0.00035859176],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034963673,0.0011690903,0.00058188284,0.0018135485,0.0007036143,0.002427586,0.0021882753,0.001213681,0.013336011],"category_scores_gemma":[0.009307502,0.00087162235,0.00096590084,0.0017736971,0.0005923395,0.003338363,0.0034574857,0.0011941235,0.003453243],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019226715,0.0020475246,0.016566789,0.0025280623,0.00046435077,0.003600467,0.009689311,0.017183352,0.05643704,0.015543587,0.10993318,0.7640836],"study_design_scores_gemma":[0.00094929495,0.0007662759,0.017624434,0.00061130757,0.0003505135,0.0035567388,0.001662215,0.41962272,0.11089236,0.016791284,0.42675075,0.00042213965],"about_ca_topic_score_codex":0.0040825307,"about_ca_topic_score_gemma":0.003952885,"teacher_disagreement_score":0.013336011,"about_ca_system_score_codex":0.00080781995,"about_ca_system_score_gemma":0.0008875905,"threshold_uncertainty_score":0.04461336},"labels":[],"label_agreement":null},{"id":"W3117079728","doi":"10.2196/31980","title":"Expressiveness of an International Semantic Standard for Wound Care: Mapping a Standardized Item Set for Leg Ulcers to the Systematized Nomenclature of Medicine–Clinical Terms","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung","keywords":"SNOMED CT; Terminology; Systematized Nomenclature of Medicine; Interoperability; Medicine; Health care; Documentation; Artificial intelligence; Computer science; World Wide Web","score_opus":0.03386969212034033,"score_gpt":0.3805468065372176,"score_spread":0.3466771144168773,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3117079728","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.60810596,0.0014775775,0.3372504,0.0025134094,0.00045964157,0.0069956253,0.009718986,0.00065772486,0.032820616],"genre_scores_gemma":[0.6948224,0.00047190898,0.28732085,0.00030318153,0.00003892767,0.0055406517,0.010455196,0.00014237166,0.0009045807],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.95900536,0.02366643,0.0074429186,0.0018036314,0.007603968,0.00047770765],"domain_scores_gemma":[0.9040935,0.060530838,0.0066402056,0.0125086075,0.015466898,0.0007600298],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03929461,0.000583358,0.00061542966,0.0068087704,0.0010747439,0.0031196258,0.0011298658,0.0009922453,0.0019555702],"category_scores_gemma":[0.13093378,0.00027227605,0.0011877883,0.005932624,0.002005823,0.0026435256,0.003806875,0.0011362989,0.00051527773],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013822644,0.00087383046,0.16505055,0.007818354,0.0005312686,0.0011225954,0.084267214,0.00904199,0.03336988,0.08947271,0.023390861,0.5836785],"study_design_scores_gemma":[0.000498767,0.0024784098,0.4679486,0.013598209,0.0014941534,0.004857391,0.073041946,0.05312807,0.0387198,0.1161207,0.22742131,0.0006927752],"about_ca_topic_score_codex":0.0029408813,"about_ca_topic_score_gemma":0.0031900227,"teacher_disagreement_score":0.03929461,"about_ca_system_score_codex":0.0029589985,"about_ca_system_score_gemma":0.006339384,"threshold_uncertainty_score":0.20781225},"labels":[],"label_agreement":null},{"id":"W3118287326","doi":"10.1007/978-3-030-66196-0_17","title":"Design of a Biochemistry Procedure-Oriented Ontology","year":2020,"lang":"en","type":"book-chapter","venue":"Communications in computer and information science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; University of Waterloo","funders":"","keywords":"Ontology; Computer science; Completeness (order theory); Subsequence; Cardinality (data modeling); Domain (mathematical analysis); Information retrieval; Decidability; Abstraction; Programming language; Theoretical computer science; Data mining; Mathematics","score_opus":0.037791606998318024,"score_gpt":0.29444910720483736,"score_spread":0.25665750020651934,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3118287326","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049392837,0.000092634895,0.9830147,0.00049559865,0.00011395036,0.0005304683,0.000979503,0.0031968688,0.0066370163],"genre_scores_gemma":[0.03546414,0.00031367698,0.95258874,0.00029349796,0.000044942288,0.00050965446,0.003752364,0.0009828325,0.0060501182],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9983725,0.0002175355,0.00024061518,0.0004249673,0.0005093517,0.00023496895],"domain_scores_gemma":[0.9986665,0.0003375057,0.00009548135,0.00027164328,0.00050209055,0.00012680674],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027333964,0.0007932751,0.00072654674,0.0030589143,0.0017318042,0.004158461,0.0020538373,0.001237331,0.008308751],"category_scores_gemma":[0.0030751342,0.0010306234,0.003178322,0.0020642183,0.0013225136,0.004259946,0.0031418398,0.0021986722,0.0028876814],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020926727,0.00047895676,0.0044977246,0.0012114804,0.00023094415,0.0012772125,0.0024276471,0.018883474,0.03274725,0.5994185,0.02053456,0.318083],"study_design_scores_gemma":[0.0001311657,0.00014405197,0.002086269,0.000587683,0.0006239551,0.0013150945,0.0014695319,0.17126381,0.046929214,0.21992251,0.555379,0.00014785829],"about_ca_topic_score_codex":0.010612903,"about_ca_topic_score_gemma":0.011222957,"teacher_disagreement_score":0.010612903,"about_ca_system_score_codex":0.0020813677,"about_ca_system_score_gemma":0.006191009,"threshold_uncertainty_score":0.027795494},"labels":[],"label_agreement":null},{"id":"W3119758793","doi":"10.1200/cci.20.00104","title":"College of American Pathologists Cancer Protocols: From Optimizing Cancer Patient Care to Facilitating Interoperable Reporting and Downstream Data Use","year":2021,"lang":"en","type":"article","venue":"JCO Clinical Cancer Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":35,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Cancer Institute","keywords":"Interoperability; Workflow; Downstream (manufacturing); Cancer; Medicine; Health care; Medical physics; Computer science; Data science; World Wide Web; Business; Database; Internal medicine; Political science","score_opus":0.18048038733967947,"score_gpt":0.4627083631819542,"score_spread":0.2822279758422747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3119758793","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02148689,0.015382786,0.3757018,0.4573223,0.004776003,0.004464047,0.005297,0.008115965,0.1074532],"genre_scores_gemma":[0.09235215,0.015814926,0.830123,0.03311779,0.0016883428,0.0031261507,0.009256416,0.0016638645,0.012857391],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.92528296,0.041077435,0.00995494,0.004252088,0.017828653,0.0016039935],"domain_scores_gemma":[0.80909795,0.065915525,0.017869623,0.035774853,0.05952514,0.0118168015],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.104265645,0.00074805116,0.0006207698,0.009824297,0.003406591,0.009620804,0.004409297,0.002105565,0.004608463],"category_scores_gemma":[0.15481997,0.0008310169,0.0008591605,0.013386115,0.0051489845,0.01287356,0.011415912,0.0064208526,0.002610638],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010109708,0.00017834155,0.016393634,0.0012863703,0.000088446606,0.00024106647,0.0055839103,0.0013793713,0.0021983876,0.077528015,0.28239745,0.6126238],"study_design_scores_gemma":[0.000053108124,0.00010264309,0.009591566,0.0023188547,0.00006150166,0.00069362065,0.0044040196,0.0026959255,0.0031748156,0.06758503,0.9091803,0.00013863842],"about_ca_topic_score_codex":0.021989007,"about_ca_topic_score_gemma":0.020428864,"teacher_disagreement_score":0.89573437,"about_ca_system_score_codex":0.009817861,"about_ca_system_score_gemma":0.0524795,"threshold_uncertainty_score":0.55141604},"labels":[],"label_agreement":null},{"id":"W3121056140","doi":"","title":"Patología digital: una perspectiva de la industria y la red de Quebec, Canadá","year":2020,"lang":"es","type":"article","venue":"I+S: Revista de la Sociedad Española de Informática y Salud","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Humanities; Political science; Art","score_opus":0.01000335566180297,"score_gpt":0.2815280198743938,"score_spread":0.2715246642125908,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3121056140","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.440155,0.010744946,0.0036829917,0.09909745,0.0006131357,0.00014660011,0.0055208164,0.00023284159,0.4398062],"genre_scores_gemma":[0.881924,0.00636782,0.0018928773,0.0035973997,0.00004996112,0.000025488165,0.00075511594,0.00007090797,0.10531647],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99886984,0.00012885692,0.000023546932,0.00011326795,0.000489109,0.0003752646],"domain_scores_gemma":[0.99773693,0.00033565884,0.00013005645,0.000081080776,0.0012390525,0.0004772499],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011025113,0.00025990157,0.00022040475,0.0019183471,0.011186984,0.0095368475,0.00087507075,0.0010717087,0.013340899],"category_scores_gemma":[0.0023217173,0.00018366527,0.0002970814,0.0054708906,0.003975712,0.002227618,0.0017707928,0.0012637561,0.0005248041],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029807448,0.00017301907,0.16080424,0.00059300615,0.0001175439,0.0037738406,0.085461274,0.0031783846,0.0035243486,0.23458509,0.19661307,0.31087816],"study_design_scores_gemma":[0.000018519155,0.000028310173,0.16272154,0.00045227684,0.000057970712,0.00026925397,0.12078558,0.001919859,0.000714135,0.007506155,0.7054428,0.00008358904],"about_ca_topic_score_codex":0.99737334,"about_ca_topic_score_gemma":0.9986664,"teacher_disagreement_score":0.09477058,"about_ca_system_score_codex":0.09477058,"about_ca_system_score_gemma":0.12278127,"threshold_uncertainty_score":0.68761194},"labels":[],"label_agreement":null},{"id":"W3121831946","doi":"10.1186/s42826-020-00068-8","title":"Establishment and application of information resource of mutant mice in RIKEN BioResource Research Center","year":2021,"lang":"en","type":"review","venue":"Laboratory Animal Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute of Genetics; Bio-oriented Technology Research Advancement Institution; Japan Society for the Promotion of Science; Ministry of Education, Culture, Sports, Science and Technology; Cabinet Office, Government of Japan; Research Organization of Information and Systems; RIKEN; Japan Agency for Medical Research and Development","keywords":"Database; Phenome; Interoperability; Resource (disambiguation); Biology; Mutant; Computational biology; Arabidopsis; Strain (injury); Research center; Phenotype; Biotechnology; Computer science; Genetics; World Wide Web; Gene; Political science","score_opus":0.06736680022130342,"score_gpt":0.4222670245628413,"score_spread":0.3549002243415379,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3121831946","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024502655,0.12890184,0.14559226,0.0020321843,0.0010061428,0.0027904906,0.5955589,0.03262092,0.066994645],"genre_scores_gemma":[0.015210419,0.047975805,0.1381365,0.00073187024,0.00015078472,0.001833556,0.7815273,0.0030852188,0.011348614],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976688,0.0002460762,0.00070165837,0.00035315062,0.00086063566,0.00016962888],"domain_scores_gemma":[0.99563605,0.00057260034,0.0007304509,0.0008360122,0.001818785,0.00040610318],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038602143,0.0013190599,0.0023553227,0.010010285,0.0007575907,0.0022953455,0.0038084688,0.0010702534,0.010731208],"category_scores_gemma":[0.005182244,0.0007330481,0.0011909992,0.011881733,0.00028550648,0.0025549608,0.0022661164,0.0013132679,0.019536272],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007947962,0.0003184425,0.0072985496,0.024997296,0.00051100383,0.0016787312,0.00107389,0.0010364214,0.17972903,0.013913031,0.2469522,0.5216967],"study_design_scores_gemma":[0.0000401769,0.00006411238,0.005364793,0.0010125499,0.00027326102,0.00063718454,0.00015684802,0.00040259882,0.027464585,0.0011084478,0.9633638,0.00011146752],"about_ca_topic_score_codex":0.0057728114,"about_ca_topic_score_gemma":0.006648273,"teacher_disagreement_score":0.010731208,"about_ca_system_score_codex":0.0012884506,"about_ca_system_score_gemma":0.0040048445,"threshold_uncertainty_score":0.03589952},"labels":[],"label_agreement":null},{"id":"W3122512721","doi":"10.2196/22976","title":"Using Machine Learning to Collect and Facilitate Remote Access to Biomedical Databases: Development of the Biomedical Database Inventory","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Spanish National Plan for Scientific and Technical Research and Innovation; European Regional Development Fund; Instituto de Salud Carlos III","keywords":"Database; Computer science; Pipeline (software); Set (abstract data type); Precision and recall; Information retrieval","score_opus":0.08440698246115125,"score_gpt":0.3617222523148145,"score_spread":0.2773152698536632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3122512721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.034576207,0.004071242,0.80824995,0.004788039,0.0003169104,0.0023896957,0.013906754,0.12619257,0.005508642],"genre_scores_gemma":[0.05364343,0.0014649082,0.918983,0.000806978,0.00010831486,0.00090826553,0.020677686,0.0014194951,0.0019879858],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99108875,0.0026043425,0.0015307104,0.0017338432,0.0027130328,0.00032929654],"domain_scores_gemma":[0.9581346,0.017100217,0.004447208,0.009761788,0.008381325,0.002174822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022922926,0.0014662552,0.0015378593,0.016514309,0.00085356645,0.005425345,0.0041005835,0.0012756682,0.0022716161],"category_scores_gemma":[0.043596067,0.0013404462,0.0015876475,0.008875851,0.0007976202,0.008134033,0.00655272,0.0033110816,0.0038943326],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041939612,0.00049993774,0.02350175,0.0018046368,0.0004085297,0.00051391666,0.0011985658,0.006542668,0.017252382,0.0058170836,0.04283148,0.8992096],"study_design_scores_gemma":[0.0004595996,0.0010610473,0.044218335,0.0033577217,0.0008274115,0.0025711274,0.0020575728,0.40975493,0.14533795,0.036004767,0.3533271,0.001022402],"about_ca_topic_score_codex":0.0041117,"about_ca_topic_score_gemma":0.0052735065,"teacher_disagreement_score":0.022922926,"about_ca_system_score_codex":0.0017528248,"about_ca_system_score_gemma":0.005014558,"threshold_uncertainty_score":0.12122941},"labels":[],"label_agreement":null},{"id":"W3123984549","doi":"10.2196/17934","title":"Hybrid Deep Learning for Medication-Related Information Extraction From Clinical Texts in French: MedExt Algorithm Development Study","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":37,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Artificial intelligence; Conditional random field; Natural language processing; Information extraction; Word embedding; Recall; Security token; Machine learning; Deep learning; F1 score; Artificial neural network; Task (project management); Recurrent neural network; Embedding","score_opus":0.022526084270900134,"score_gpt":0.3437960057843904,"score_spread":0.32126992151349026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3123984549","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63205206,0.005931585,0.33532554,0.0011895617,0.0002030792,0.00072771695,0.0026039095,0.014821859,0.0071447655],"genre_scores_gemma":[0.63385075,0.001318224,0.34948623,0.0006551838,0.00007144955,0.00047797876,0.007044508,0.00032714187,0.0067684893],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99904543,0.00038551484,0.00008511761,0.0002635183,0.000119951525,0.00010046076],"domain_scores_gemma":[0.99747145,0.0017104199,0.0000857895,0.00016925279,0.0005021868,0.000061021667],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023777697,0.0014774667,0.0007670635,0.0015677678,0.00039631972,0.0010447089,0.0012168342,0.0015326695,0.0031267668],"category_scores_gemma":[0.0045684073,0.00034691923,0.00087387674,0.0011374902,0.00029691405,0.0013033901,0.0007292489,0.0011363378,0.0007661019],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067467196,0.00096599077,0.009327275,0.00046157595,0.00040777252,0.00038147444,0.00024000452,0.21203871,0.00629952,0.0017449905,0.008135106,0.7593229],"study_design_scores_gemma":[0.00009759207,0.00020532626,0.0019186235,0.000038541046,0.000078739584,0.00010002127,0.00008027513,0.98849034,0.0056764893,0.00073913543,0.0025609804,0.000013862001],"about_ca_topic_score_codex":0.02155907,"about_ca_topic_score_gemma":0.021089299,"teacher_disagreement_score":0.02155907,"about_ca_system_score_codex":0.0017425092,"about_ca_system_score_gemma":0.0019090227,"threshold_uncertainty_score":0.042867124},"labels":[],"label_agreement":null},{"id":"W3124781663","doi":"10.2196/25530","title":"Similarity-Based Unsupervised Spelling Correction Using BioWordVec: Development and Usability Study of Bacterial Culture and Antimicrobial Susceptibility Reports","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Korea Health Industry Development Institute","keywords":"Spelling; Computer science; Artificial intelligence; Natural language processing; Edit distance; Similarity (geometry); Ranking (information retrieval); Vocabulary; Pattern recognition (psychology); Speech recognition; Information retrieval; Linguistics","score_opus":0.022311373905479798,"score_gpt":0.295768969183719,"score_spread":0.2734575952782392,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3124781663","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95898926,0.0005814771,0.03630092,0.00016223946,0.00007609415,0.0009576772,0.0005392316,0.0015135459,0.0008795249],"genre_scores_gemma":[0.8183197,0.0008773353,0.17236578,0.00014249116,0.00004383784,0.0015443488,0.004034509,0.00052253035,0.0021495176],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9900064,0.0058342507,0.001350107,0.001103328,0.0014540822,0.00025180116],"domain_scores_gemma":[0.9485882,0.032991428,0.0023433366,0.0034282696,0.011639071,0.0010096959],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013031004,0.0013083073,0.0007464467,0.0028510196,0.0004537736,0.001362945,0.0014860166,0.0007456995,0.00071266264],"category_scores_gemma":[0.042814966,0.00052669266,0.0012248748,0.0015743903,0.00073702954,0.0029146248,0.0019166041,0.00083942595,0.00047782378],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002064118,0.0055324407,0.07942496,0.005863015,0.00076884864,0.0017683264,0.02095694,0.00994985,0.05405661,0.0015099436,0.009849511,0.80825555],"study_design_scores_gemma":[0.0012800213,0.026740866,0.30588463,0.0017414382,0.002172103,0.0062883478,0.021589637,0.45148504,0.14634494,0.0026953898,0.032864366,0.0009133568],"about_ca_topic_score_codex":0.0037777256,"about_ca_topic_score_gemma":0.0048569543,"teacher_disagreement_score":0.013031004,"about_ca_system_score_codex":0.00082864665,"about_ca_system_score_gemma":0.0012583194,"threshold_uncertainty_score":0.06891537},"labels":[],"label_agreement":null},{"id":"W3126565418","doi":"10.2196/21679","title":"Lexicon Development for COVID-19-related Concepts Using Open-source Word Embedding Sources: An Intrinsic and Extrinsic Evaluation","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; National Institutes of Health; Perelman School of Medicine, University of Pennsylvania; University of Pennsylvania","keywords":"Coronavirus disease 2019 (COVID-19); Lexicon; Computer science; Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Word embedding; 2019-20 coronavirus outbreak; Word (group theory); Open source; Natural language processing; Artificial intelligence; Embedding; Linguistics; Medicine; Virology; Software; Pathology; Programming language","score_opus":0.07589123149888627,"score_gpt":0.4143356395403869,"score_spread":0.3384444080415006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3126565418","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40280837,0.0041664615,0.48840052,0.0032725409,0.0016102296,0.0060913186,0.042943913,0.031122385,0.019584313],"genre_scores_gemma":[0.34631655,0.0012641008,0.5619928,0.00053577236,0.00018075929,0.003351929,0.08081139,0.0023633006,0.0031834578],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98643005,0.0057485076,0.002474438,0.0017789348,0.0031567516,0.00041132967],"domain_scores_gemma":[0.93338937,0.045621824,0.0029289613,0.005228843,0.011800305,0.0010307648],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015416133,0.0024488473,0.0010262027,0.00981826,0.001319795,0.0051334244,0.0016518302,0.0018359203,0.0059807496],"category_scores_gemma":[0.082423754,0.00062129745,0.0021650118,0.004791855,0.0011069092,0.009554804,0.008212369,0.0021410834,0.0035736184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025257827,0.0014351858,0.06646337,0.008420877,0.0010235645,0.0017355864,0.0076253586,0.012019621,0.030299313,0.015455548,0.08656084,0.7664351],"study_design_scores_gemma":[0.0012889261,0.0021985485,0.0592045,0.0032720747,0.0016804449,0.004318805,0.013382773,0.62198275,0.06410609,0.039418835,0.18836613,0.00078012387],"about_ca_topic_score_codex":0.0040798387,"about_ca_topic_score_gemma":0.005620226,"teacher_disagreement_score":0.015416133,"about_ca_system_score_codex":0.0018461066,"about_ca_system_score_gemma":0.0030989014,"threshold_uncertainty_score":0.08152926},"labels":[],"label_agreement":null},{"id":"W3128776966","doi":"10.21742/ajnnia.2020.4.1.01","title":"Detecting Drug-Drug Interaction (DDI) over the Social Media using Convolution Neural Network Deep Learning","year":2020,"lang":"en","type":"article","venue":"Asia-Pacific Journal of Neural Networks and Its Applications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Word embedding; Computer science; Artificial intelligence; Convolutional neural network; Classifier (UML); Machine learning; Feature vector; Feature learning; Support vector machine; Artificial neural network; SemEval; Representation (politics); Word (group theory); Natural language processing; Deep learning; Embedding; Task (project management)","score_opus":0.02376257876376233,"score_gpt":0.2792338724243422,"score_spread":0.2554712936605799,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3128776966","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7543781,0.009006028,0.19836311,0.0036230816,0.0003882579,0.00035227535,0.016751792,0.0048095495,0.012327937],"genre_scores_gemma":[0.9348253,0.0021552516,0.04958781,0.00051890296,0.00017647193,0.00011104563,0.007993579,0.00005286262,0.004578772],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99950767,0.00009399646,0.000051845458,0.00013112293,0.00014545217,0.00006988815],"domain_scores_gemma":[0.9985978,0.0007041144,0.0003597539,0.00010155031,0.00018032841,0.000056407473],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005850448,0.0008495811,0.000565966,0.0027731233,0.0002590488,0.00069482543,0.00040820206,0.0006415426,0.0012457756],"category_scores_gemma":[0.0024033298,0.00020376862,0.0005994176,0.0017188584,0.00027483096,0.0015089299,0.0007020881,0.0006992692,0.0007177619],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009223892,0.0010968422,0.09823772,0.0008459958,0.0006631629,0.0014148164,0.0002948186,0.04622869,0.02729203,0.002579878,0.018140508,0.8022831],"study_design_scores_gemma":[0.000030676892,0.00031436572,0.039830863,0.00011631056,0.00017622214,0.00083804963,0.0002333104,0.9192803,0.019525338,0.009474586,0.010131508,0.0000485419],"about_ca_topic_score_codex":0.0061318786,"about_ca_topic_score_gemma":0.011043205,"teacher_disagreement_score":0.0061318786,"about_ca_system_score_codex":0.00078496506,"about_ca_system_score_gemma":0.00051236263,"threshold_uncertainty_score":0.0121923685},"labels":[],"label_agreement":null},{"id":"W3129472010","doi":"10.26577/iam.2020.v1.i1.01","title":"Advices On Search And Critical Appraisal Of Biomedical Literature Part I, General Workflow","year":2020,"lang":"en","type":"article","venue":"Interdisciplinary Approaches to Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Workflow; Critical appraisal; Subject (documents); Computer science; Systematic review; Data science; Management science; Grey literature; Engineering ethics; MEDLINE; Knowledge management; Medicine; Alternative medicine; World Wide Web; Political science; Engineering; Pathology","score_opus":0.09836991778909843,"score_gpt":0.3577134191607525,"score_spread":0.2593435013716541,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3129472010","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00631232,0.04492554,0.4472445,0.30395073,0.040208913,0.057706974,0.008716437,0.02708177,0.063852854],"genre_scores_gemma":[0.013807133,0.027283872,0.86175966,0.03143979,0.012894259,0.028639082,0.0025064668,0.0023170542,0.019352697],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.80582976,0.11957017,0.05141488,0.0042923084,0.016935792,0.0019570747],"domain_scores_gemma":[0.35923493,0.41889897,0.041148588,0.029422222,0.13921538,0.012079902],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.17804094,0.0032064694,0.0040817033,0.02181845,0.0038841257,0.00845301,0.005152537,0.0074376324,0.054284383],"category_scores_gemma":[0.4505648,0.0025936041,0.0047017136,0.011117971,0.0051874463,0.0095335515,0.0068403045,0.0075467667,0.042165656],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006750203,0.00018876704,0.00130542,0.017595196,0.00023683431,0.0012288337,0.009853831,0.00047820958,0.003442999,0.006715661,0.5990107,0.35926855],"study_design_scores_gemma":[0.0006986763,0.00023266343,0.0040393225,0.033898417,0.00029832526,0.0025140692,0.005586975,0.0023134123,0.0026190504,0.04944447,0.8978278,0.0005268217],"about_ca_topic_score_codex":0.003567499,"about_ca_topic_score_gemma":0.006342722,"teacher_disagreement_score":0.8219591,"about_ca_system_score_codex":0.005331545,"about_ca_system_score_gemma":0.038513813,"threshold_uncertainty_score":0.9415817},"labels":[],"label_agreement":null},{"id":"W3131377851","doi":"10.1101/2021.02.18.430807","title":"Relating simulation studies by provenance—Developing a family of Wnt signaling models","year":2021,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"Deutsche Forschungsgemeinschaft","keywords":"Computer science; Metadata; Ontology; Reuse; Variety (cybernetics); Simulation modeling; Exploit; Data modeling; Data science; Key (lock); Software engineering; World Wide Web; Engineering; Artificial intelligence","score_opus":0.04067446743548804,"score_gpt":0.2783287480595395,"score_spread":0.23765428062405147,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3131377851","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019337064,0.0003300531,0.9703916,0.0015486489,0.00007950427,0.00034036633,0.0007932512,0.0014848177,0.005694866],"genre_scores_gemma":[0.20262612,0.0008135573,0.7896192,0.00035483757,0.00006155785,0.00044615247,0.0029801803,0.00079350075,0.002304926],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98536706,0.006945158,0.0015541172,0.0014073618,0.0042592804,0.00046706514],"domain_scores_gemma":[0.9527062,0.023779944,0.0035174198,0.014103661,0.0048541557,0.0010385403],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01948189,0.00084919546,0.0007036978,0.0043123816,0.0017295305,0.0046564955,0.0024009906,0.0017855706,0.0029226616],"category_scores_gemma":[0.0503648,0.0010424113,0.0038271768,0.0033125265,0.0032964242,0.009088708,0.0073499964,0.0034424819,0.0005979944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033175282,0.00033525843,0.031175105,0.0007455948,0.00023124459,0.0020388162,0.005356929,0.18381533,0.008146372,0.6795112,0.006837223,0.08147516],"study_design_scores_gemma":[0.0000604357,0.00010676404,0.002412618,0.00046570707,0.00011169745,0.000847108,0.0011464213,0.49318475,0.014261903,0.38399825,0.10329481,0.00010965534],"about_ca_topic_score_codex":0.008602183,"about_ca_topic_score_gemma":0.0071777687,"teacher_disagreement_score":0.01948189,"about_ca_system_score_codex":0.0028824252,"about_ca_system_score_gemma":0.0040893997,"threshold_uncertainty_score":0.10303128},"labels":[],"label_agreement":null},{"id":"W3133630228","doi":"10.3390/asi4010021","title":"An Ontological Approach for Early Detection of Suspected COVID-19 among COPD Patients","year":2021,"lang":"en","type":"article","venue":"Applied System Innovation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"COPD; Context (archaeology); Medicine; Intensive care medicine; Coronavirus disease 2019 (COVID-19); Pandemic; Pulmonary disease; Computer science; Disease; Internal medicine; Infectious disease (medical specialty)","score_opus":0.023216048108430604,"score_gpt":0.2734746940607538,"score_spread":0.2502586459523232,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3133630228","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12372666,0.0007764867,0.85571444,0.0043417444,0.00017540375,0.00084917375,0.0021751414,0.002625492,0.009615395],"genre_scores_gemma":[0.4909991,0.00066635816,0.5025791,0.0006320018,0.000059974976,0.00029608858,0.0028071294,0.00007944667,0.0018808002],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997988,0.00050688576,0.00034710087,0.00040586357,0.0005987938,0.00015338107],"domain_scores_gemma":[0.9974279,0.0010239093,0.0003415592,0.00035758794,0.0006862668,0.00016271592],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022057581,0.0005851318,0.0004449908,0.0043657366,0.0013411925,0.002141717,0.0010260659,0.0010518393,0.0008430937],"category_scores_gemma":[0.0069663124,0.00029725168,0.0014260922,0.0017928461,0.00065189914,0.0027868436,0.002262091,0.0009060226,0.00032176368],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00065332296,0.0014535926,0.13448785,0.0012994183,0.0005136214,0.0037170064,0.006836471,0.03430473,0.05582774,0.088669725,0.015035802,0.6572007],"study_design_scores_gemma":[0.00007357545,0.00021838327,0.051042337,0.0005671204,0.0009345189,0.002524248,0.0045771,0.7545715,0.029153699,0.090167336,0.06594983,0.00022043867],"about_ca_topic_score_codex":0.013320421,"about_ca_topic_score_gemma":0.0145276785,"teacher_disagreement_score":0.013320421,"about_ca_system_score_codex":0.0016248797,"about_ca_system_score_gemma":0.0037497662,"threshold_uncertainty_score":0.026485741},"labels":[],"label_agreement":null},{"id":"W3134842774","doi":"10.5281/zenodo.4650697","title":"FIDEO: Food Interactions with Drugs Evidence Ontology","year":2020,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Hospitalier de l’Université de Montréal","funders":"Horizon 2020 Framework Programme; Agence Nationale de la Recherche; European Commission","keywords":"Ontology; Computer science; Data science; Epistemology; Philosophy","score_opus":0.028004965911583243,"score_gpt":0.26799095604291795,"score_spread":0.2399859901313347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3134842774","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0116438,0.0050833295,0.85547435,0.007674383,0.0011445977,0.0012404079,0.06377341,0.010472502,0.043493245],"genre_scores_gemma":[0.07724564,0.006536429,0.8100138,0.002969424,0.0005259087,0.0014344306,0.087369904,0.0013941422,0.012510348],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974407,0.00049656467,0.0004905212,0.0004484219,0.00094649947,0.0001773948],"domain_scores_gemma":[0.99623364,0.0018244147,0.00041922677,0.00057160197,0.00068289845,0.00026823528],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026581478,0.0010055922,0.0008878916,0.006486874,0.0014337486,0.0031621845,0.0018000742,0.0019778598,0.0061274865],"category_scores_gemma":[0.008097012,0.0006392971,0.0023545488,0.004135329,0.0011081304,0.005644338,0.0028827651,0.0019195342,0.0025759584],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000430267,0.00041966944,0.008871271,0.00545606,0.00046456806,0.0024803195,0.0017969764,0.010239415,0.010019877,0.46256393,0.121036,0.37622163],"study_design_scores_gemma":[0.00005228134,0.000042572177,0.0034511068,0.00068017206,0.00017652426,0.0014704668,0.00037075515,0.011242755,0.003064199,0.07380461,0.90555924,0.000085313535],"about_ca_topic_score_codex":0.012041247,"about_ca_topic_score_gemma":0.011699232,"teacher_disagreement_score":0.012041247,"about_ca_system_score_codex":0.0025881391,"about_ca_system_score_gemma":0.0056868023,"threshold_uncertainty_score":0.023942292},"labels":[],"label_agreement":null},{"id":"W3135406142","doi":"10.1101/2021.03.10.382333","title":"Capturing scientific knowledge in computable form","year":2021,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Army Research Office; Defense Advanced Research Projects Agency; National Institutes of Health","keywords":"Computer science; Data science; Context (archaeology); Suite; Pace; Ontology; World Wide Web; Sociology of scientific knowledge; Knowledge management","score_opus":0.017729289672007268,"score_gpt":0.24310698633088573,"score_spread":0.22537769665887847,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3135406142","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019612294,0.0009479694,0.9412623,0.0027653244,0.00022546524,0.00019693447,0.0084974505,0.003295669,0.023196472],"genre_scores_gemma":[0.28238347,0.0021049383,0.6881663,0.00045689713,0.00031007742,0.0005428836,0.018773628,0.0006830945,0.0065785944],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9946601,0.0013620572,0.00072690035,0.0009667108,0.0020367322,0.00024758154],"domain_scores_gemma":[0.9867578,0.007242064,0.000914234,0.0031715687,0.0016328641,0.0002814632],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0030190684,0.00094360445,0.001005378,0.007447608,0.0011380825,0.010992162,0.0016579183,0.0012315748,0.008812042],"category_scores_gemma":[0.032387685,0.00065760355,0.0017421286,0.009691666,0.0032184676,0.011496439,0.005643484,0.0016939814,0.0027588014],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013660302,0.00010064323,0.0033809554,0.0011881996,0.00012274033,0.0009047686,0.0014384937,0.05674966,0.002485907,0.7811233,0.01466204,0.1377067],"study_design_scores_gemma":[0.000022462253,0.00001620179,0.0004909247,0.00028917854,0.000052234067,0.00016522485,0.0003621732,0.09956133,0.002502188,0.83520305,0.06130641,0.000028638127],"about_ca_topic_score_codex":0.0048850067,"about_ca_topic_score_gemma":0.0044483924,"teacher_disagreement_score":0.9969809,"about_ca_system_score_codex":0.002177269,"about_ca_system_score_gemma":0.0026569928,"threshold_uncertainty_score":0.029479206},"labels":[],"label_agreement":null},{"id":"W3136401565","doi":"10.1101/2021.03.12.21253461","title":"Biomedical Discovery through the integrative Biomedical Knowledge Hub (iBKH)","year":2021,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; HEC Montréal","funders":"","keywords":"Biomedicine; Computer science; Repurposing; Scalability; Data science; Knowledge extraction; Pipeline (software); Knowledge graph; Artificial intelligence; Bioinformatics; Engineering","score_opus":0.02807745064568505,"score_gpt":0.3187101302495178,"score_spread":0.2906326796038327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3136401565","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.067845054,0.0049069794,0.84692246,0.003718182,0.0003543812,0.000564131,0.031002272,0.030917754,0.013768693],"genre_scores_gemma":[0.34383816,0.002350462,0.6127055,0.0005571744,0.00017184927,0.0002626051,0.035538573,0.0013696218,0.0032059196],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987213,0.00033643315,0.00008375881,0.00039638826,0.00036103543,0.00010104301],"domain_scores_gemma":[0.9967428,0.0013562933,0.00037361437,0.00079435995,0.0004811535,0.00025190052],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030612666,0.00066209614,0.00082379935,0.008305699,0.0009680802,0.0027445585,0.0013827196,0.0007372027,0.00468531],"category_scores_gemma":[0.0072355033,0.000496933,0.0012881286,0.0063172816,0.0010791642,0.003414416,0.004567386,0.0011961488,0.0016646952],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010401254,0.00036357026,0.022990767,0.005107943,0.0013137575,0.0014255524,0.0018681629,0.08741124,0.036140658,0.24748287,0.09547472,0.49938077],"study_design_scores_gemma":[0.00018128828,0.00014562964,0.010266876,0.0006454783,0.00068998354,0.00075612176,0.000687308,0.34035042,0.032944992,0.4185715,0.19461992,0.00014046411],"about_ca_topic_score_codex":0.0045373156,"about_ca_topic_score_gemma":0.0054963306,"teacher_disagreement_score":0.008305699,"about_ca_system_score_codex":0.001117517,"about_ca_system_score_gemma":0.0026680087,"threshold_uncertainty_score":0.016189694},"labels":[],"label_agreement":null},{"id":"W3150158741","doi":"","title":"Clinical documents and their parts","year":2020,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Computer science; Information retrieval; Natural language processing","score_opus":0.02760939050140008,"score_gpt":0.2893533990914993,"score_spread":0.2617440085900992,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150158741","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019316614,0.03845902,0.12514669,0.044507787,0.01878585,0.0015845664,0.45124134,0.013674884,0.28728324],"genre_scores_gemma":[0.108563736,0.035821185,0.12696485,0.0064184163,0.008249749,0.0009291129,0.4069547,0.0044156145,0.3016826],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973667,0.00047999495,0.00055247976,0.00053171144,0.0009294515,0.00013969942],"domain_scores_gemma":[0.9921234,0.0032771933,0.0006471637,0.0014799131,0.0018421276,0.00063014944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020878878,0.0010520407,0.0011156448,0.007961426,0.001096068,0.0072248825,0.000758357,0.0016306163,0.11383661],"category_scores_gemma":[0.01550578,0.00063169666,0.0007891994,0.014068567,0.0011805968,0.0042737382,0.0024899037,0.001534912,0.075640514],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004976583,0.0000775779,0.0014690543,0.0020537875,0.000057463494,0.0009916079,0.0005510906,0.0009641821,0.010086425,0.079786286,0.50987196,0.3935929],"study_design_scores_gemma":[0.000035632984,0.000023494287,0.0016878606,0.00030144554,0.00003391971,0.0010183618,0.0002333976,0.00079100137,0.0024709213,0.02450497,0.96887577,0.000023079185],"about_ca_topic_score_codex":0.0017151127,"about_ca_topic_score_gemma":0.0008240874,"teacher_disagreement_score":0.11383661,"about_ca_system_score_codex":0.0012076084,"about_ca_system_score_gemma":0.003104283,"threshold_uncertainty_score":0.3808214},"labels":[],"label_agreement":null},{"id":"W3150833277","doi":"10.5220/0010375500002865","title":"Applying PySCMGroup to Breast Cancer Biomarkers Discovery","year":2021,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; Université Laval","funders":"","keywords":"Breast cancer; Computer science; Cancer; Computational biology; Medicine; Internal medicine; Biology","score_opus":0.012618409430692035,"score_gpt":0.2802179238282263,"score_spread":0.26759951439753427,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3150833277","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18649761,0.0047733174,0.5995608,0.003119407,0.0007734973,0.001547075,0.0917438,0.09984419,0.012140282],"genre_scores_gemma":[0.24933338,0.0010692803,0.6164803,0.00060110434,0.0002129017,0.0006533992,0.1238357,0.002302116,0.0055117505],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998092,0.00044963008,0.00024980426,0.00045592003,0.0006005531,0.00015203515],"domain_scores_gemma":[0.99682415,0.0012380866,0.00021107755,0.0009882948,0.0005228925,0.00021553299],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002743996,0.00093489845,0.00090413267,0.00755378,0.001256025,0.0017824542,0.0010643019,0.0008443137,0.0041801375],"category_scores_gemma":[0.009152024,0.0002973455,0.0020571502,0.0046977475,0.0004822881,0.001826515,0.0034036103,0.00071765145,0.0021851426],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012813245,0.0005896968,0.036449928,0.0022195864,0.0015517104,0.0016518409,0.0009507531,0.011039385,0.024808574,0.00945333,0.06444154,0.8455623],"study_design_scores_gemma":[0.0004626077,0.000863367,0.04965175,0.00084677374,0.0017934031,0.0029619248,0.0018260389,0.3821626,0.06350428,0.1090305,0.38671002,0.00018664803],"about_ca_topic_score_codex":0.006032107,"about_ca_topic_score_gemma":0.01120395,"teacher_disagreement_score":0.00755378,"about_ca_system_score_codex":0.0006975806,"about_ca_system_score_gemma":0.0037789498,"threshold_uncertainty_score":0.014511764},"labels":[],"label_agreement":null},{"id":"W3154324864","doi":"10.1002/1873-3468.14067","title":"Sharing biological data: why, when, and how","year":2021,"lang":"en","type":"editorial","venue":"FEBS Letters","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":73,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Vector Institute; University of Toronto; Princess Margaret Cancer Centre; University Health Network","funders":"Canadian Institutes of Health Research; National Institutes of Health; Natural Sciences and Engineering Research Council of Canada; Broad Institute; Vlaamse regering; Fonds Wetenschappelijk Onderzoek","keywords":"Metadata; Computer science; Data sharing; Data element; USable; Reuse; Data discovery; Interoperability; Data mapping; Documentation; Data quality; Data access; Data type; Data dictionary; Transparency (behavior); World Wide Web; Data science; Database; Service (business)","score_opus":0.039709001670996974,"score_gpt":0.28221444865483014,"score_spread":0.24250544698383317,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154324864","genre_codex":"commentary","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005142199,0.034759082,0.26708704,0.6468549,0.009267336,0.0012022207,0.0010893176,0.0017567939,0.032841116],"genre_scores_gemma":[0.14818998,0.06172589,0.5585787,0.1881835,0.012075118,0.0041065305,0.0035723369,0.0041976525,0.019370258],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.84254205,0.08966409,0.010665903,0.015951665,0.03489279,0.006283588],"domain_scores_gemma":[0.8011599,0.10154267,0.00950605,0.04877494,0.02731021,0.011706169],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.15429972,0.0018272585,0.002772621,0.0049018827,0.010142583,0.04187405,0.0077926354,0.0132027725,0.009361845],"category_scores_gemma":[0.20482326,0.0021751865,0.002743002,0.0063475654,0.051858082,0.08072422,0.024633598,0.016716322,0.01051405],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016154729,0.00019251402,0.007575072,0.004481636,0.00035592247,0.00043016442,0.019691296,0.0013462666,0.0024611733,0.48809984,0.18071428,0.2944902],"study_design_scores_gemma":[0.00003706548,0.00006488856,0.0011611162,0.006316062,0.00008309202,0.0004551888,0.009366656,0.0008868589,0.001648195,0.59133637,0.3884704,0.00017404821],"about_ca_topic_score_codex":0.009769897,"about_ca_topic_score_gemma":0.0067167766,"teacher_disagreement_score":0.99220735,"about_ca_system_score_codex":0.008865093,"about_ca_system_score_gemma":0.027977707,"threshold_uncertainty_score":0.81602466},"labels":[],"label_agreement":null},{"id":"W3154325449","doi":"10.17504/protocols.io.bid7ka9n","title":"SARS-CoV-2 NCBI consensus submission protocol: GenBank v2","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Protocol (science); GenBank; Biology; Metadata; Coronavirus disease 2019 (COVID-19); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Computer science; World Wide Web; Computational biology; Medicine; Genetics; Pathology","score_opus":0.08101030844248443,"score_gpt":0.368760272244654,"score_spread":0.2877499638021696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154325449","genre_codex":"dataset","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060232556,0.002754355,0.09176075,0.0064919335,0.0064727264,0.019183071,0.74737316,0.028388213,0.091552556],"genre_scores_gemma":[0.007236409,0.0015153526,0.056198463,0.0043665655,0.00072476984,0.018720886,0.8604468,0.0068225283,0.043968193],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98717624,0.0044802516,0.0031048243,0.0014702043,0.0027267896,0.0010417976],"domain_scores_gemma":[0.9816296,0.0029565794,0.0013571824,0.0038958627,0.00916297,0.0009977799],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011415821,0.0018267044,0.0020541823,0.0048293667,0.0041189524,0.0041068955,0.004263274,0.0042269505,0.24774839],"category_scores_gemma":[0.026199818,0.0022284375,0.0011673538,0.0052503943,0.001267236,0.0037818772,0.004047645,0.004239318,0.3364258],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067965314,0.0001218401,0.0006208227,0.002002088,0.000023300587,0.00021969277,0.00037575845,0.0001736854,0.01674321,0.0033856307,0.953181,0.02247323],"study_design_scores_gemma":[0.00021572773,0.00023536358,0.0016909112,0.001042317,0.000029757695,0.00034872748,0.000392412,0.00033678944,0.009931821,0.0027313659,0.982951,0.00009388385],"about_ca_topic_score_codex":0.0037913204,"about_ca_topic_score_gemma":0.0045354106,"teacher_disagreement_score":0.98858416,"about_ca_system_score_codex":0.0018018744,"about_ca_system_score_gemma":0.008047822,"threshold_uncertainty_score":0.8288009},"labels":[],"label_agreement":null},{"id":"W3154688461","doi":"10.17504/protocols.io.bsypnfvn","title":"SARS-CoV-2 NCBI submission workflow + guidance for structuring and releasing metadata v1","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Metadata; Workflow; GenBank; Computer science; Protocol (science); Process (computing); World Wide Web; Biology; Database","score_opus":0.04934370549416287,"score_gpt":0.32647289021690723,"score_spread":0.27712918472274434,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3154688461","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035968197,0.0010101764,0.2972592,0.0063969754,0.0031361373,0.014720457,0.3555191,0.23905046,0.079310656],"genre_scores_gemma":[0.008300959,0.0016659664,0.35796592,0.005374929,0.00088369567,0.020235278,0.4394846,0.08907299,0.07701569],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9876935,0.0031205772,0.002739335,0.001393221,0.0040281205,0.0010252481],"domain_scores_gemma":[0.9676605,0.009223635,0.001907265,0.009533207,0.0099597825,0.001715623],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026930105,0.0020818987,0.0017596588,0.006276599,0.0034240251,0.009007982,0.003872778,0.003474471,0.21184143],"category_scores_gemma":[0.0617455,0.0032303804,0.0019643502,0.004640024,0.0012037869,0.006278259,0.009160053,0.004053164,0.3612253],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056082633,0.00011241799,0.0011112036,0.0015030978,0.000030105262,0.00029780887,0.00069272955,0.00040212885,0.0061927564,0.007155983,0.909698,0.07224291],"study_design_scores_gemma":[0.00009113539,0.00006653304,0.001032109,0.0008342072,0.000013055145,0.00022190964,0.0002795786,0.0005473276,0.0056426073,0.0046311445,0.9865465,0.00009387477],"about_ca_topic_score_codex":0.0072843037,"about_ca_topic_score_gemma":0.0074416357,"teacher_disagreement_score":0.21184143,"about_ca_system_score_codex":0.0030493203,"about_ca_system_score_gemma":0.01037377,"threshold_uncertainty_score":0.70868015},"labels":[],"label_agreement":null},{"id":"W3155548469","doi":"10.2196/28247","title":"A Novel Metric to Quantify the Effect of Pathway Enrichment Evaluation With Respect to Biomedical Text-Mined Terms: Development and Feasibility Study","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Metric (unit); Computer science; Robustness (evolution); Data mining; Drug discovery; Inference; Computational biology; Machine learning; Artificial intelligence; Bioinformatics; Biology; Gene","score_opus":0.036136888950801034,"score_gpt":0.3561194037950782,"score_spread":0.3199825148442772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3155548469","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.216428,0.0032982659,0.7730619,0.0005241844,0.00019625585,0.00059286266,0.002155197,0.0013382701,0.0024051147],"genre_scores_gemma":[0.59527445,0.0005201204,0.4016607,0.00010238238,0.00006915037,0.0004827812,0.0013526725,0.00010059401,0.00043723176],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99068505,0.0031292397,0.0011241749,0.0011911278,0.0036109206,0.00025945264],"domain_scores_gemma":[0.9401098,0.04445922,0.0041881804,0.003131146,0.007389561,0.00072213047],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010967694,0.0015700168,0.0014134377,0.0058851345,0.0005470424,0.0017688392,0.0012218236,0.0013697611,0.0012024805],"category_scores_gemma":[0.055479594,0.00025065092,0.0011769627,0.0043691783,0.0010089427,0.0026352534,0.001498605,0.0010856795,0.00026203063],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020149802,0.00069581746,0.10048803,0.0016775945,0.0009719603,0.0004817292,0.0002977148,0.22101398,0.031962056,0.010619441,0.002732849,0.62704384],"study_design_scores_gemma":[0.00013268407,0.0026906782,0.03708458,0.00013334725,0.0002943298,0.0011937564,0.0002230424,0.9127101,0.031569783,0.010067365,0.00372686,0.00017349188],"about_ca_topic_score_codex":0.0020973454,"about_ca_topic_score_gemma":0.001807857,"teacher_disagreement_score":0.010967694,"about_ca_system_score_codex":0.0014523842,"about_ca_system_score_gemma":0.0016852989,"threshold_uncertainty_score":0.058003366},"labels":[],"label_agreement":null},{"id":"W3156105266","doi":"10.20416/lsrsps.v8i2.6","title":"Une ontologie dispositionnelle du risque","year":2021,"lang":"fr","type":"article","venue":"Lato Sensu Revue de la Société de philosophie des sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Humanities; Philosophy; Political science","score_opus":0.0382790820874199,"score_gpt":0.3195180736343917,"score_spread":0.28123899154697185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3156105266","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029121235,0.0028094323,0.8690821,0.006312361,0.0008068064,0.0005842031,0.004031701,0.0013984552,0.08585373],"genre_scores_gemma":[0.33474353,0.0077717416,0.5931468,0.0023512747,0.0009315537,0.001718275,0.009142881,0.0009389711,0.049254928],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9934814,0.0018116249,0.00081068603,0.0016184793,0.0019833634,0.00029448987],"domain_scores_gemma":[0.99048436,0.004007207,0.0007097913,0.001489924,0.002912961,0.00039585904],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046879724,0.0011283488,0.0006008699,0.005024167,0.0028734775,0.0076068197,0.0010254087,0.0016255783,0.0077262763],"category_scores_gemma":[0.014394711,0.0008788933,0.0022392375,0.0034446418,0.0045916163,0.01329639,0.0035423907,0.003738923,0.0018901008],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007837281,0.00006524258,0.0045298194,0.0005608688,0.00011333833,0.00072560157,0.004759392,0.0016443258,0.0033013674,0.9106489,0.0069547296,0.06661804],"study_design_scores_gemma":[0.00003116598,0.000081328806,0.008136337,0.0005461499,0.00015205657,0.0019124661,0.0030600422,0.0098692225,0.0036729644,0.6200047,0.35240862,0.0001250199],"about_ca_topic_score_codex":0.012691433,"about_ca_topic_score_gemma":0.007880826,"teacher_disagreement_score":0.012691433,"about_ca_system_score_codex":0.0030674022,"about_ca_system_score_gemma":0.0041621015,"threshold_uncertainty_score":0.025846958},"labels":[],"label_agreement":null},{"id":"W3156270706","doi":"10.1007/s12021-021-09522-x","title":"Correction to: A Standards Organization for Open and FAIR Neuroscience: the International Neuroinformatics Coordinating Facility","year":2021,"lang":"en","type":"erratum","venue":"Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Baycrest Hospital; University of Toronto; Montreal Neurological Institute and Hospital; McGill University","funders":"","keywords":"Neuroinformatics; Open science; Computer science; Neuroscience; Data science; World Wide Web; Cognitive science; Psychology","score_opus":0.019964500308350504,"score_gpt":0.30392556032655227,"score_spread":0.28396106001820176,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3156270706","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00013399309,0.0003860486,0.0010810472,0.106612824,0.8813055,0.000047301262,0.0018228559,0.00054366037,0.008066788],"genre_scores_gemma":[0.008023565,0.0029634312,0.00843863,0.18416019,0.25369146,0.00043500063,0.0048282016,0.0030175843,0.53444195],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9921084,0.0010774954,0.0012634647,0.0009327095,0.0040571294,0.0005607946],"domain_scores_gemma":[0.9274974,0.0131793395,0.0022423856,0.0042189234,0.050258312,0.002603634],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0065016267,0.0016908307,0.001776854,0.004412528,0.004925169,0.0073388685,0.0031282164,0.009821018,0.06439057],"category_scores_gemma":[0.0954631,0.0010460368,0.0012785469,0.0028393096,0.0033215075,0.0035327072,0.0028092593,0.016464878,0.05803927],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008003208,0.0000027256176,0.000022011438,0.000020710613,0.000002127634,0.00005710552,0.000012447533,0.000014495048,0.0000110530145,0.00056984153,0.997675,0.0016044327],"study_design_scores_gemma":[0.000018779634,0.0000062698336,0.00035185044,0.00018248869,0.000013747573,0.00016703803,0.00008202067,0.00017013968,0.00016175945,0.0016710266,0.9971493,0.00002559926],"about_ca_topic_score_codex":0.050499436,"about_ca_topic_score_gemma":0.05773213,"teacher_disagreement_score":0.99687177,"about_ca_system_score_codex":0.006587927,"about_ca_system_score_gemma":0.012236727,"threshold_uncertainty_score":0.21540791},"labels":[],"label_agreement":null},{"id":"W3157054249","doi":"","title":"A BERT-based Model for Drug-Drug Interaction Extraction from Drug Labels.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug; Computer science; Extraction (chemistry); Drug-drug interaction; Pharmacology; Chromatography; Chemistry; Medicine","score_opus":0.009371832549338802,"score_gpt":0.27967689358487835,"score_spread":0.27030506103553953,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157054249","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020560548,0.0037047458,0.9336427,0.0027168863,0.0002930825,0.0007480424,0.020625824,0.00975427,0.007953835],"genre_scores_gemma":[0.3264592,0.0023846088,0.6227697,0.00123204,0.00020032341,0.0014772762,0.03538312,0.00044632918,0.009647438],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99920684,0.00015982817,0.00009651579,0.0001800537,0.00028954312,0.00006721017],"domain_scores_gemma":[0.99832207,0.000986367,0.00009523365,0.0001306555,0.00038620865,0.00007948347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011495501,0.0010166444,0.0010159482,0.003828435,0.00073180423,0.0013722426,0.0019956373,0.001440329,0.0043089557],"category_scores_gemma":[0.004650541,0.0004200896,0.0019838721,0.0033861126,0.00044288527,0.0020101853,0.0014153604,0.0011635466,0.002484508],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013562131,0.0005589052,0.011296496,0.0017444546,0.00065935595,0.0011758506,0.0003893266,0.27425262,0.0068427115,0.054594662,0.062323213,0.5848062],"study_design_scores_gemma":[0.0000457391,0.000082083425,0.0015404938,0.00010734332,0.00015086355,0.00026019494,0.00006350008,0.92504513,0.0018090812,0.05006654,0.020801557,0.000027475604],"about_ca_topic_score_codex":0.025414454,"about_ca_topic_score_gemma":0.038726546,"teacher_disagreement_score":0.025414454,"about_ca_system_score_codex":0.001624032,"about_ca_system_score_gemma":0.0026607458,"threshold_uncertainty_score":0.050533056},"labels":[],"label_agreement":null},{"id":"W3157414336","doi":"10.25073/2588-1086/vnucsce.237","title":"Single Concatenated Input is Better than Indenpendent Multiple-input for CNNs to Predict Chemical-induced Disease Relation from Literature","year":2020,"lang":"en","type":"article","venue":"VNU Journal of Science Computer Science and Communication Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"National Foundation for Science and Technology Development","keywords":"Concatenation (mathematics); Computer science; Biomedical text mining; Convolutional neural network; Benchmark (surveying); Relation (database); Named-entity recognition; Artificial intelligence; Natural language processing; Data mining; Text mining; Mathematics; Task (project management); Cartography","score_opus":0.02200202434629777,"score_gpt":0.2493643346008226,"score_spread":0.22736231025452483,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3157414336","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67569035,0.013060476,0.26201865,0.00392723,0.0017225852,0.00035577436,0.004402859,0.017617963,0.02120399],"genre_scores_gemma":[0.91241723,0.0017696584,0.06573629,0.00088253315,0.00020473564,0.00011249486,0.0053709513,0.00013460714,0.013371466],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996828,0.00003445453,0.000033677294,0.000113139824,0.000062939675,0.000072968905],"domain_scores_gemma":[0.99947804,0.00019853744,0.00006176566,0.0000669255,0.00015227345,0.000042559837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00079297007,0.001717105,0.00063711684,0.0006956957,0.00033834984,0.0010143631,0.0013611919,0.0010454486,0.0034003898],"category_scores_gemma":[0.0020199008,0.00032111554,0.00065389625,0.0006335763,0.00036365684,0.0017737284,0.000882722,0.0013937339,0.0017101722],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010462741,0.0007179362,0.015097485,0.00043165637,0.00038387883,0.00045937882,0.00013859331,0.09441406,0.034399327,0.0016517368,0.0140256295,0.8372341],"study_design_scores_gemma":[0.000028458251,0.0003075171,0.0025490334,0.00005641864,0.00014599324,0.000094439114,0.00004183766,0.96988523,0.021856407,0.0016523723,0.003358663,0.00002366266],"about_ca_topic_score_codex":0.011721632,"about_ca_topic_score_gemma":0.020110657,"teacher_disagreement_score":0.011721632,"about_ca_system_score_codex":0.00097623083,"about_ca_system_score_gemma":0.0010799828,"threshold_uncertainty_score":0.023306847},"labels":[],"label_agreement":null},{"id":"W3158507570","doi":"","title":"A Hybrid Model for Drug-Drug Interaction Extraction from Structured Product Labeling Documents.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug; Computer science; Extraction (chemistry); Product (mathematics); Chemistry; Pharmacology; Chromatography; Mathematics; Medicine","score_opus":0.00845630159188394,"score_gpt":0.28219377991580724,"score_spread":0.2737374783239233,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158507570","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02990404,0.0026934824,0.9209554,0.0018365327,0.0002289604,0.0007922657,0.022713233,0.014916262,0.0059597166],"genre_scores_gemma":[0.21200621,0.0014494065,0.74517304,0.00071668695,0.000097997756,0.0011254201,0.032133657,0.00035357114,0.006943987],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999338,0.00013321055,0.000098307355,0.00015483184,0.00023360076,0.00004212273],"domain_scores_gemma":[0.99854434,0.00082708395,0.000081058985,0.00015177674,0.00034060288,0.000055252232],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001061505,0.0006538825,0.0007022708,0.0038793655,0.00053456874,0.0015331251,0.0014440446,0.0011434244,0.0025332188],"category_scores_gemma":[0.0032639336,0.0003148312,0.0015176241,0.00345155,0.00029202862,0.0019268484,0.0010019338,0.0006835034,0.001767671],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011605889,0.00070831465,0.0121556325,0.0019709277,0.0008220857,0.0012074352,0.0005761519,0.12662052,0.011543143,0.035834983,0.05912767,0.7482726],"study_design_scores_gemma":[0.000089646375,0.00013430693,0.002733468,0.00017313061,0.00033552517,0.00045843757,0.00015541144,0.907385,0.0058363727,0.040868502,0.0417825,0.000047619684],"about_ca_topic_score_codex":0.014537232,"about_ca_topic_score_gemma":0.030908123,"teacher_disagreement_score":0.014537232,"about_ca_system_score_codex":0.0010399885,"about_ca_system_score_gemma":0.0021234076,"threshold_uncertainty_score":0.028905272},"labels":[],"label_agreement":null},{"id":"W3158833611","doi":"10.1101/2021.05.04.21256134","title":"Automated Medical Chart Review for Breast Cancer: A Novel Natural Language Processing Software System","year":2021,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Workflow; Pipeline (software); Computer science; Context (archaeology); Software; Chart; Breast cancer; Health care; Artificial intelligence; Data science; Medicine; Cancer; Database; Statistics","score_opus":0.01634403388543198,"score_gpt":0.321052070614684,"score_spread":0.304708036729252,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158833611","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09156505,0.001038848,0.60403836,0.0030559876,0.00025249773,0.0015467154,0.029198429,0.2626258,0.006678348],"genre_scores_gemma":[0.2019333,0.000493147,0.757808,0.00091983925,0.00015760369,0.0006011216,0.031799223,0.0016794242,0.0046082777],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983537,0.00032247568,0.00022358315,0.00057942246,0.00046231202,0.00005846404],"domain_scores_gemma":[0.9967981,0.0015396407,0.00030701866,0.00046505049,0.0007073801,0.00018276833],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019456189,0.00060378393,0.0005122508,0.0021859424,0.00053471274,0.0015337967,0.0011326445,0.0006986051,0.0041209613],"category_scores_gemma":[0.007082861,0.000387661,0.00063719123,0.0014038734,0.00031285704,0.0015417432,0.0013308858,0.0006229948,0.0026761617],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010052929,0.00073398463,0.01797755,0.0012508794,0.0002466463,0.0019617288,0.0010972057,0.008213617,0.07745725,0.0058243265,0.16121937,0.72301215],"study_design_scores_gemma":[0.00080765586,0.00045285947,0.034195613,0.0003601755,0.00026929038,0.0038254852,0.00074054895,0.64193445,0.103281826,0.019474803,0.19438756,0.00026974533],"about_ca_topic_score_codex":0.0072095185,"about_ca_topic_score_gemma":0.008948483,"teacher_disagreement_score":0.0072095185,"about_ca_system_score_codex":0.0009646542,"about_ca_system_score_gemma":0.002904482,"threshold_uncertainty_score":0.014335096},"labels":[],"label_agreement":null},{"id":"W3158858160","doi":"","title":"Overview of the TAC 2019 Track on Drug-Drug Interaction Extraction from Drug Labels.","year":2019,"lang":"en","type":"article","venue":"Theory and applications of categories","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug; Drug-drug interaction; Track (disk drive); Computer science; Extraction (chemistry); Pharmacology; Medicine; Chemistry; Chromatography","score_opus":0.010892914647400673,"score_gpt":0.285857184977048,"score_spread":0.2749642703296473,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3158858160","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008261792,0.01992224,0.23331583,0.0045327023,0.0017587955,0.0016573629,0.5248839,0.12916826,0.076499045],"genre_scores_gemma":[0.008658357,0.004045862,0.14747632,0.0017905578,0.00026709054,0.0007602986,0.8198925,0.004015337,0.013093613],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9959347,0.0006469805,0.00056308764,0.0008052408,0.0017923429,0.000257572],"domain_scores_gemma":[0.992157,0.0023554647,0.00047759197,0.0018180186,0.002649735,0.00054221606],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005312374,0.0023162554,0.0023359747,0.016825087,0.0019860344,0.005753601,0.0043453723,0.0026312247,0.03125702],"category_scores_gemma":[0.013658262,0.0009799233,0.0030712322,0.013714991,0.00060504547,0.0058241053,0.0036726855,0.0024750023,0.031020503],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004323626,0.00021277097,0.002185167,0.0026321558,0.00045016987,0.00020864962,0.00013666622,0.0023497902,0.0053311298,0.011147263,0.75912374,0.21579017],"study_design_scores_gemma":[0.00017257936,0.00017697466,0.004160817,0.00077036844,0.00031135554,0.0003438113,0.00008130418,0.01673437,0.0068130526,0.012153417,0.95817935,0.00010262307],"about_ca_topic_score_codex":0.029021576,"about_ca_topic_score_gemma":0.037612546,"teacher_disagreement_score":0.03125702,"about_ca_system_score_codex":0.0029370205,"about_ca_system_score_gemma":0.007428049,"threshold_uncertainty_score":0.104565084},"labels":[],"label_agreement":null},{"id":"W3159272427","doi":"10.1002/ca.23741","title":"General histological woes: Definition and classification of tissues","year":2021,"lang":"en","type":"review","venue":"Clinical Anatomy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Acadia University; Dalhousie University","funders":"","keywords":"Connective tissue; Pathology; Medicine; Human body; Cell type; Extracellular matrix; Function (biology); Anatomy; Biology; Computational biology; Evolutionary biology; Cell biology; Cell; Genetics","score_opus":0.1985433242739811,"score_gpt":0.45351699492834163,"score_spread":0.2549736706543605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159272427","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013176593,0.92163527,0.04697561,0.005700823,0.0033767743,0.00036276004,0.0018221256,0.000550239,0.018258777],"genre_scores_gemma":[0.008164311,0.9118447,0.056068864,0.0041149976,0.001540239,0.0005305102,0.005734313,0.00014156838,0.011860602],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983937,0.00032470655,0.0003843406,0.00030732542,0.0005015852,0.000088405875],"domain_scores_gemma":[0.9982975,0.0005328928,0.00027217605,0.00016982581,0.00063343754,0.00009418343],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038260443,0.001697408,0.0019236588,0.013048463,0.0006283359,0.0030786244,0.0035755443,0.0015266292,0.0034677102],"category_scores_gemma":[0.0036615543,0.00066709955,0.001334793,0.00917663,0.0033207561,0.0067774244,0.001750721,0.0028733136,0.0041044992],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005885795,0.000059872982,0.0010890628,0.01846665,0.00012068025,0.00030404818,0.000608492,0.00058994594,0.0030028063,0.063672855,0.07011282,0.8419138],"study_design_scores_gemma":[0.000009111829,0.000028219712,0.002319749,0.005464277,0.00008846521,0.0014809789,0.00038076637,0.0002597421,0.00081548747,0.02032333,0.9687942,0.00003569552],"about_ca_topic_score_codex":0.004873805,"about_ca_topic_score_gemma":0.0044769174,"teacher_disagreement_score":0.013048463,"about_ca_system_score_codex":0.0027591332,"about_ca_system_score_gemma":0.0043940926,"threshold_uncertainty_score":0.020234287},"labels":[],"label_agreement":null},{"id":"W3159409148","doi":"10.1093/database/baab021","title":"Increasing metadata coverage of SRA BioSample entries using deep learning–based named entity recognition","year":2021,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; National Institutes of Health; National Institute of General Medical Sciences; Canadian Institute for Advanced Research","keywords":"Metadata; Computer science; Named-entity recognition; Information retrieval; Artificial intelligence; Deep learning; Natural language processing; World Wide Web","score_opus":0.040499717845373796,"score_gpt":0.29040079590886664,"score_spread":0.24990107806349285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3159409148","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7660538,0.0017028359,0.14144291,0.0009786191,0.00030667768,0.00017630671,0.05096965,0.031771813,0.00659751],"genre_scores_gemma":[0.6188905,0.00054789276,0.23124553,0.00064094935,0.00009824872,0.00032377042,0.14315642,0.00092333945,0.004173242],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992398,0.00010312361,0.00007126324,0.00035090008,0.00016685668,0.00006802417],"domain_scores_gemma":[0.9977678,0.00079824793,0.00029866575,0.0004402101,0.000597556,0.00009746701],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0016393727,0.0008448887,0.00057585083,0.0020175185,0.0005313392,0.00089185516,0.00094720337,0.0006944997,0.0014222872],"category_scores_gemma":[0.0047716135,0.00026257156,0.00088306353,0.0013985683,0.00033723513,0.0018735273,0.0013267178,0.0010109749,0.0019574342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012904335,0.00085281907,0.1779504,0.0009631697,0.0004879589,0.00088519795,0.00094143767,0.05530151,0.17998274,0.0024709424,0.053553745,0.52531976],"study_design_scores_gemma":[0.000067215224,0.00036413717,0.058262005,0.00018696126,0.00022129354,0.0003728359,0.0006635878,0.7698882,0.13170463,0.0057821567,0.032358523,0.00012852416],"about_ca_topic_score_codex":0.0065763067,"about_ca_topic_score_gemma":0.014799806,"teacher_disagreement_score":0.99836063,"about_ca_system_score_codex":0.0005611786,"about_ca_system_score_gemma":0.0005910556,"threshold_uncertainty_score":0.013076067},"labels":[],"label_agreement":null},{"id":"W3161027543","doi":"10.2196/29667","title":"A Word Pair Dataset for Semantic Similarity and Relatedness in Korean Medical Vocabulary: Reference Development and Validation","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Korea Health Industry Development Institute","keywords":"Natural language processing; Computer science; Similarity (geometry); Vocabulary; Word (group theory); Word embedding; Task (project management); Artificial intelligence; Semantic similarity; Set (abstract data type); Health informatics; Information retrieval; Unified Medical Language System; Embedding; Linguistics; Medicine; Pathology","score_opus":0.03191483108306381,"score_gpt":0.31704466072013704,"score_spread":0.28512982963707323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3161027543","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6162636,0.004060528,0.046592094,0.0008009635,0.00071373134,0.008496229,0.3081327,0.005854665,0.009085399],"genre_scores_gemma":[0.19132718,0.00057272194,0.072327316,0.00027253188,0.00009414983,0.0077043884,0.7251674,0.00031261132,0.0022217021],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99244654,0.002204836,0.002304182,0.0015279927,0.0012414663,0.00027505634],"domain_scores_gemma":[0.98312116,0.005960446,0.0015698636,0.0039445953,0.004556676,0.00084734126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0072523854,0.0017238188,0.000996526,0.009116691,0.0016989827,0.0014862154,0.0025884877,0.0020239777,0.004649746],"category_scores_gemma":[0.024378607,0.00045134142,0.0017659824,0.004675097,0.0013072292,0.0034175003,0.0046247984,0.0018613207,0.006120738],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0032802203,0.0048370655,0.15905194,0.013519652,0.0013972421,0.0049002855,0.0082776705,0.008662685,0.035636365,0.0072165676,0.2838804,0.4693399],"study_design_scores_gemma":[0.0032525123,0.003767702,0.44192785,0.0022071705,0.0013665174,0.009025086,0.012671874,0.08588516,0.042809676,0.007293947,0.3889674,0.00082520506],"about_ca_topic_score_codex":0.005567947,"about_ca_topic_score_gemma":0.008453181,"teacher_disagreement_score":0.009116691,"about_ca_system_score_codex":0.0012059474,"about_ca_system_score_gemma":0.0024165716,"threshold_uncertainty_score":0.038354695},"labels":[],"label_agreement":null},{"id":"W3164061191","doi":"10.3233/shti210267","title":"Building a Knowledge Graph Representing Causal Associations Between Risk Factors and Incidence of Breast Cancer","year":2021,"lang":"en","type":"book-chapter","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Graph; Breast cancer; Traverse; Semantic Web; Data science; Node (physics); Information retrieval; Medicine; Cancer; Theoretical computer science","score_opus":0.04695940951190088,"score_gpt":0.37357860831075945,"score_spread":0.3266191987988586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3164061191","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017405417,0.0013821982,0.95333636,0.0015212644,0.000098068536,0.00020043793,0.005019782,0.0022500472,0.018786307],"genre_scores_gemma":[0.08417029,0.0021876316,0.90098834,0.00034552737,0.000038359794,0.0001821027,0.0074561452,0.00020069732,0.004430862],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995933,0.00012215973,0.00004027029,0.00012590355,0.000102241305,0.000016167685],"domain_scores_gemma":[0.9988213,0.0008715716,0.000077237404,0.00008424885,0.00012370429,0.000021884263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007877143,0.00055992114,0.0003237993,0.004486395,0.0005473882,0.0015058416,0.0009284264,0.00069999695,0.0057473877],"category_scores_gemma":[0.0027820729,0.00029606943,0.0009999078,0.003398624,0.00044748306,0.0026478537,0.0009421504,0.0005209265,0.0013323558],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000115250856,0.00013289996,0.0068186526,0.002294574,0.00022216032,0.001516432,0.001559339,0.068106264,0.0120106,0.2493521,0.022882488,0.6349892],"study_design_scores_gemma":[0.00003961288,0.000089791465,0.0047380044,0.0008737558,0.00042245237,0.0015648765,0.0011289672,0.24675626,0.014873646,0.47212663,0.25730142,0.00008459916],"about_ca_topic_score_codex":0.004352138,"about_ca_topic_score_gemma":0.0068699014,"teacher_disagreement_score":0.0057473877,"about_ca_system_score_codex":0.0008004828,"about_ca_system_score_gemma":0.0012050414,"threshold_uncertainty_score":0.019226909},"labels":[],"label_agreement":null},{"id":"W3164732864","doi":"10.3233/shti210187","title":"A Knowledge Graph of Mechanistic Associations Between COVID-19, Diabetes Mellitus and Kidney Diseases","year":2021,"lang":"en","type":"book-chapter","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Coronavirus disease 2019 (COVID-19); Knowledge graph; Computer science; Disease; Data science; Graph; Diabetes mellitus; Medicine; Infectious disease (medical specialty); Artificial intelligence; Theoretical computer science; Pathology","score_opus":0.0509293438551032,"score_gpt":0.35655161490071136,"score_spread":0.30562227104560813,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3164732864","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022593886,0.010985363,0.836251,0.012083271,0.00068447867,0.00044268396,0.019644722,0.0035556117,0.09375904],"genre_scores_gemma":[0.12653711,0.011334584,0.81906044,0.001844946,0.0002951892,0.00031225133,0.017575031,0.0002417837,0.022798667],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995721,0.00011206499,0.00004346232,0.00013643542,0.000118946526,0.000016932632],"domain_scores_gemma":[0.9981634,0.0013839271,0.00013680694,0.000085830354,0.0001848415,0.0000451932],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00076974253,0.0006216111,0.00035493932,0.00531244,0.0008248024,0.0019751075,0.00091250194,0.0007965504,0.010583183],"category_scores_gemma":[0.0031581991,0.0003050591,0.0010559718,0.004651806,0.00070161256,0.0031493036,0.00095046865,0.0008065777,0.0014871965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000920337,0.00010915314,0.0057947887,0.002005737,0.00025747996,0.002488493,0.0011065293,0.025952306,0.0056526633,0.47329828,0.0639707,0.41927192],"study_design_scores_gemma":[0.000025779364,0.00005739876,0.005063829,0.0009477803,0.00037136461,0.002353513,0.000689805,0.06611743,0.003423582,0.42682874,0.49406144,0.000059370253],"about_ca_topic_score_codex":0.0070161,"about_ca_topic_score_gemma":0.009040304,"teacher_disagreement_score":0.010583183,"about_ca_system_score_codex":0.0011426286,"about_ca_system_score_gemma":0.001639084,"threshold_uncertainty_score":0.035404325},"labels":[],"label_agreement":null},{"id":"W3164966020","doi":"10.3233/shti210182","title":"OCRx: Canadian Drug Ontology","year":2021,"lang":"en","type":"book-chapter","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université Laval; Université de Montréal; Centre Hospitalier de l’Université de Montréal","funders":"","keywords":"Ontology; Interoperability; Usability; Computer science; World Wide Web; Identification (biology); Product (mathematics); Drug; Medicine; Pharmacology","score_opus":0.035010465255414146,"score_gpt":0.32484636725794813,"score_spread":0.289835902002534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3164966020","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025230525,0.011038323,0.092796974,0.010490687,0.0016725279,0.00043471713,0.03338694,0.008232,0.8394247],"genre_scores_gemma":[0.02341156,0.01807202,0.21380597,0.0044366345,0.00034054965,0.0003431893,0.057442456,0.005371088,0.67677647],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986364,0.00007570132,0.000049025173,0.0001479279,0.0009641862,0.0001267287],"domain_scores_gemma":[0.99866056,0.00019950792,0.00004137278,0.000113127644,0.0008442957,0.00014110551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012010701,0.0010227306,0.00041133663,0.0064529125,0.0027380511,0.0046731024,0.001851975,0.0012777774,0.042114757],"category_scores_gemma":[0.0022180018,0.00047494326,0.00055824924,0.009275473,0.0017836194,0.004053468,0.0019168922,0.0016534178,0.014470336],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019727715,0.00001631021,0.00024434525,0.0003693682,0.000006222704,0.00015087148,0.00069857866,0.0010020687,0.0015658116,0.27189845,0.586422,0.13760623],"study_design_scores_gemma":[9.096322e-7,7.23446e-7,0.00007581801,0.000030091038,0.0000011779748,0.00004533038,0.000031721123,0.00017205939,0.00015017195,0.0015636042,0.9979236,0.00000479634],"about_ca_topic_score_codex":0.73623943,"about_ca_topic_score_gemma":0.826562,"teacher_disagreement_score":0.73623943,"about_ca_system_score_codex":0.030456228,"about_ca_system_score_gemma":0.033807445,"threshold_uncertainty_score":0.5306278},"labels":[],"label_agreement":null},{"id":"W3165450499","doi":"10.3233/shti210284","title":"Facilitating Study and Item Level Browsing for Clinical and Epidemiological COVID-19 Studies","year":2021,"lang":"en","type":"book-chapter","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University Health Centre","funders":"Klaus Tschira Stiftung; Deutsche Forschungsgemeinschaft","keywords":"Metadata; Computer science; World Wide Web; Coronavirus disease 2019 (COVID-19); Information retrieval; Data science; Medicine; Pathology","score_opus":0.44330333808630284,"score_gpt":0.5242817043907728,"score_spread":0.08097836630446997,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3165450499","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009742449,0.0077313622,0.71806526,0.008866794,0.00111321,0.001779887,0.02169095,0.043294065,0.18771596],"genre_scores_gemma":[0.023463285,0.0051118825,0.83926916,0.0029860432,0.0005082592,0.0010404255,0.01794581,0.008708527,0.100966536],"study_design_codex":"design_other","study_design_gemma":"design_other","domain_scores_codex":[0.9978034,0.0009148907,0.0003162233,0.00024912832,0.0006279574,0.00008835991],"domain_scores_gemma":[0.98589516,0.010549175,0.00039848062,0.0015414942,0.000930957,0.00068469567],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.007515306,0.0008332931,0.0008158195,0.006214562,0.0008554641,0.007773214,0.0016100951,0.001446056,0.046531703],"category_scores_gemma":[0.013116141,0.0007737435,0.0009837095,0.006159054,0.00080682704,0.008763794,0.005496474,0.0016813403,0.031513333],"study_design_candidate":"design_other","study_design_consensus":"design_other","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016018766,0.000104357045,0.003180454,0.0023570233,0.00007357014,0.00058805314,0.0031538517,0.00096343947,0.00761333,0.0823422,0.30953112,0.5899324],"study_design_scores_gemma":[0.000038718495,0.000037365266,0.002272845,0.0010209719,0.00003756678,0.0010042953,0.00073655625,0.0033718932,0.0025601687,0.04141915,0.9474367,0.00006385842],"about_ca_topic_score_codex":0.0016591303,"about_ca_topic_score_gemma":0.0049578934,"teacher_disagreement_score":0.9922268,"about_ca_system_score_codex":0.0013604473,"about_ca_system_score_gemma":0.0024223293,"threshold_uncertainty_score":0.15566409},"labels":[],"label_agreement":null},{"id":"W3167010744","doi":"10.1007/978-3-030-77211-6_13","title":"Semantic Web Framework to Computerize Staged Reflex Testing Protocols to Mitigate Underutilization of Pathology Tests for Diagnosing Pituitary Disorders","year":2021,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nova Scotia Health Authority; Dalhousie University","funders":"","keywords":"Computer science; Reflex; Protocol (science); Presentation (obstetrics); Psychological intervention; Test (biology); Intervention (counseling); Medical physics; Medicine; Intensive care medicine; Pathology; Surgery; Psychiatry; Alternative medicine","score_opus":0.0447250111411927,"score_gpt":0.3374563926055723,"score_spread":0.2927313814643796,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3167010744","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026317192,0.00014353982,0.9673872,0.00031925034,0.00010285002,0.00027797592,0.0018811717,0.023520749,0.0037356839],"genre_scores_gemma":[0.06837473,0.00038350446,0.91476667,0.00058143906,0.000058217116,0.00053648977,0.008472606,0.0021214616,0.0047048577],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986872,0.00024257602,0.00019713843,0.00024181102,0.0005196651,0.000111588204],"domain_scores_gemma":[0.9982943,0.0006290475,0.00010588465,0.0003917103,0.00045124744,0.00012780899],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029485533,0.0008047148,0.0007472154,0.0024520084,0.00088453636,0.002814554,0.0019896217,0.0013316206,0.004790402],"category_scores_gemma":[0.0041539688,0.00047426234,0.0022287015,0.0014358718,0.0007111711,0.0033499936,0.0026154649,0.0016023177,0.0022232123],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070208043,0.0010969087,0.0043556904,0.0011104569,0.0004409775,0.0014029647,0.000913592,0.083968066,0.018880162,0.2935352,0.074651964,0.51894194],"study_design_scores_gemma":[0.00009584304,0.00007498239,0.0010522249,0.00028543553,0.00022127983,0.0005353777,0.00025849996,0.594287,0.023520065,0.19571793,0.18385197,0.00009931629],"about_ca_topic_score_codex":0.008579115,"about_ca_topic_score_gemma":0.011739858,"teacher_disagreement_score":0.008579115,"about_ca_system_score_codex":0.0013430013,"about_ca_system_score_gemma":0.0029630414,"threshold_uncertainty_score":0.017058313},"labels":[],"label_agreement":null},{"id":"W3167370031","doi":"10.1093/jamiaopen/ooab035","title":"Aligning an interface terminology to the Logical Observation Identifiers Names and Codes (LOINC®)","year":2021,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Centre Hospitalier de l’Université de Montréal","funders":"","keywords":"Identifier; Computer science; Interface (matter); Artificial intelligence; Terminology; Interoperability; Information retrieval; Programming language; World Wide Web","score_opus":0.05609294620949171,"score_gpt":0.35243087355985114,"score_spread":0.2963379273503594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3167370031","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24193108,0.0014679894,0.71996695,0.0011993046,0.0005503651,0.00087873853,0.008569471,0.008687962,0.016748097],"genre_scores_gemma":[0.412366,0.000575045,0.55776787,0.00042699784,0.000118544864,0.0007465697,0.021839835,0.0025233582,0.003635875],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9903267,0.0035830943,0.0020893717,0.001753317,0.001800034,0.0004474639],"domain_scores_gemma":[0.98237926,0.0063755163,0.0029458166,0.0033818753,0.004484089,0.00043338563],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009408455,0.0009270002,0.00062861474,0.00740471,0.0009475942,0.0041533415,0.0012223107,0.0010482855,0.0026954855],"category_scores_gemma":[0.02021856,0.00039686565,0.00093671243,0.0054989797,0.0012817218,0.0031059699,0.002549794,0.0010961472,0.0017855045],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019212865,0.00045306244,0.17353688,0.0053139916,0.0004018714,0.0017818647,0.02626704,0.012130675,0.1344302,0.09323593,0.030259801,0.5202674],"study_design_scores_gemma":[0.00015771798,0.000837826,0.17511268,0.0018192467,0.0007563485,0.002916716,0.013768844,0.06376173,0.15204027,0.029227117,0.55920184,0.00039967452],"about_ca_topic_score_codex":0.009733411,"about_ca_topic_score_gemma":0.0043165567,"teacher_disagreement_score":0.009733411,"about_ca_system_score_codex":0.002978806,"about_ca_system_score_gemma":0.005072872,"threshold_uncertainty_score":0.0497573},"labels":[],"label_agreement":null},{"id":"W3170771812","doi":"10.1093/database/baab069","title":"OBO Foundry in 2021: operationalizing open data principles to evaluate ontologies","year":2021,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":206,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Human Genome Research Institute; National Institutes of Health","keywords":"Computer science; Operationalization; World Wide Web; Data science","score_opus":0.20312183623828936,"score_gpt":0.4157198848331276,"score_spread":0.21259804859483825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3170771812","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.066896185,0.00060236367,0.82981217,0.015964601,0.0015451371,0.004429367,0.00895909,0.023754694,0.04803633],"genre_scores_gemma":[0.15025926,0.00027434833,0.81475496,0.0021071844,0.00017445472,0.0021337671,0.018640442,0.0050770133,0.006578508],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.8808183,0.055065706,0.014028128,0.008171415,0.036778785,0.0051376563],"domain_scores_gemma":[0.7500297,0.071702175,0.0142851975,0.076058626,0.080444194,0.0074800663],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.1664337,0.0015754693,0.0021855016,0.009912846,0.005139089,0.023840865,0.005191683,0.003453957,0.0045147403],"category_scores_gemma":[0.23832104,0.0015126322,0.0026370236,0.005486563,0.008647444,0.021849371,0.020314638,0.0064283097,0.002579347],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011592534,0.0010290712,0.033302665,0.0010029402,0.00027595038,0.0006103182,0.00978562,0.023028875,0.008816028,0.58681285,0.07928936,0.2548871],"study_design_scores_gemma":[0.0002992352,0.000578221,0.009434195,0.0020681322,0.00015521262,0.00039152292,0.0053740684,0.118314326,0.028204838,0.4501958,0.38450563,0.00047891852],"about_ca_topic_score_codex":0.024880331,"about_ca_topic_score_gemma":0.018325223,"teacher_disagreement_score":0.9948083,"about_ca_system_score_codex":0.011909521,"about_ca_system_score_gemma":0.023264822,"threshold_uncertainty_score":0.8801961},"labels":[],"label_agreement":null},{"id":"W3171450272","doi":"10.1136/rapm-2020-102451","title":"Standardizing nomenclature in regional anesthesia: an ASRA-ESRA Delphi consensus study of abdominal wall, paraspinal, and chest wall blocks","year":2021,"lang":"en","type":"article","venue":"Regional Anesthesia & Pain Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":264,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"St. Thomas Hospital; Toronto Western Hospital; Dalhousie University; University of Ottawa; University of Toronto","funders":"","keywords":"CLARITY; Nomenclature; Delphi method; Harmonization; Standardization; Consensus conference; Medicine; Delphi; Medical physics; Computer science; Political science; Artificial intelligence; Taxonomy (biology); Law; Library science","score_opus":0.03398223843376127,"score_gpt":0.29789204853143125,"score_spread":0.26390981009767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3171450272","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.862556,0.0032211752,0.07851882,0.0073349862,0.0007661063,0.030443542,0.00043302472,0.00014736374,0.016578896],"genre_scores_gemma":[0.91342896,0.0012635442,0.056237094,0.0017308592,0.000083757644,0.02597329,0.00034536576,0.000068021014,0.0008691806],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.53254664,0.39287668,0.03444445,0.00696215,0.02444581,0.008724205],"domain_scores_gemma":[0.66879535,0.21440446,0.0149019305,0.010097467,0.085665904,0.006134852],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.40331963,0.0011176468,0.0014589563,0.00514761,0.0046362597,0.0037786628,0.0031487006,0.0029056503,0.0023354848],"category_scores_gemma":[0.32046285,0.0013005692,0.002078364,0.0037288847,0.0054352293,0.0055334354,0.012769107,0.0031502978,0.0006097055],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015706158,0.0009936901,0.04991837,0.010872932,0.0004725847,0.0007533231,0.7123501,0.0032495763,0.005120609,0.014385832,0.00884087,0.1914715],"study_design_scores_gemma":[0.0012232658,0.0049573346,0.08433414,0.017127708,0.0008266602,0.001391514,0.7829302,0.017662568,0.008121968,0.01920674,0.061353512,0.00086438976],"about_ca_topic_score_codex":0.0035642646,"about_ca_topic_score_gemma":0.0032334942,"teacher_disagreement_score":0.40331963,"about_ca_system_score_codex":0.00788238,"about_ca_system_score_gemma":0.026476037,"threshold_uncertainty_score":0.73581314},"labels":[],"label_agreement":null},{"id":"W3172474051","doi":"10.14288/1.0396097","title":"Context-informed population health knowledge translation : a case study","year":2021,"lang":"en","type":"article","venue":"cIRcle (University of British Columbia)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Context (archaeology); Population; Knowledge translation; Computer science; Medicine; Knowledge management; Geography; Environmental health","score_opus":0.03093323622244469,"score_gpt":0.26232845895744705,"score_spread":0.23139522273500235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3172474051","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89608973,0.0025325532,0.02815505,0.016854972,0.00038270105,0.0024817814,0.00020975877,0.00009494443,0.053198393],"genre_scores_gemma":[0.95880437,0.0027555972,0.025756342,0.00302626,0.0000901643,0.001374133,0.00012531834,0.00007568313,0.007992149],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.96862656,0.026542345,0.00062677916,0.0009869569,0.0014646054,0.0017527443],"domain_scores_gemma":[0.9680434,0.024892533,0.0014573248,0.0018353361,0.0014834611,0.0022880412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022034124,0.00091966544,0.00081105594,0.0023877136,0.01940364,0.007845412,0.0043068225,0.0074912165,0.0051678754],"category_scores_gemma":[0.028957734,0.00089297286,0.001025079,0.0039902553,0.0107816225,0.0071336506,0.0131491665,0.0061245426,0.0010351741],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016189252,0.0013959908,0.007857629,0.0007680264,0.00003269045,0.048702393,0.87321657,0.00061820337,0.00073269306,0.017873386,0.0040201945,0.0446204],"study_design_scores_gemma":[0.00007158471,0.0004499968,0.0026899332,0.0011155722,0.000053407082,0.015086128,0.9010114,0.0015247008,0.0013241853,0.009088157,0.06752243,0.00006250009],"about_ca_topic_score_codex":0.0063957404,"about_ca_topic_score_gemma":0.013933732,"teacher_disagreement_score":0.022034124,"about_ca_system_score_codex":0.006866727,"about_ca_system_score_gemma":0.0098181255,"threshold_uncertainty_score":0.11652899},"labels":[],"label_agreement":null},{"id":"W3176752316","doi":"10.1093/database/baab035","title":"Which methods are the most effective in enabling novice users to participate in ontology creation? A usability study","year":2021,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Government of Canada; Agriculture and Agri-Food Canada; University of Manitoba","funders":"Canadian Institutes of Health Research; National Science Foundation","keywords":"Usability; Wizard; Computer science; Ontology; Interoperability; World Wide Web; Set (abstract data type); Think aloud protocol; Data curation; Data science; Information retrieval; Human–computer interaction","score_opus":0.03992296776808297,"score_gpt":0.40249750443283333,"score_spread":0.36257453666475037,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176752316","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9641719,0.0019655684,0.025899585,0.000875129,0.000091461356,0.0017752115,0.00023430456,0.0008941712,0.004092736],"genre_scores_gemma":[0.88873184,0.0013783606,0.10452113,0.0005424241,0.000059461436,0.0026806179,0.00035206415,0.00043911612,0.0012949957],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9609003,0.025766708,0.00517904,0.0028189803,0.004163184,0.0011718156],"domain_scores_gemma":[0.630496,0.3292925,0.009905518,0.009320769,0.018300751,0.0026843941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.049793772,0.001453377,0.0012832681,0.00274489,0.0014058748,0.004731631,0.0011797897,0.0021468715,0.0015513285],"category_scores_gemma":[0.18332829,0.00096932845,0.001206735,0.0015210327,0.0014216974,0.005027238,0.0020677124,0.0009613841,0.0005276149],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009032154,0.0054921205,0.14908518,0.014006149,0.0009158808,0.00093727617,0.09004281,0.0017331031,0.035839852,0.0014574755,0.0061892155,0.68526876],"study_design_scores_gemma":[0.0057352125,0.054193586,0.52140737,0.008442292,0.004615962,0.004356417,0.1614803,0.055073056,0.09001995,0.011460832,0.08078992,0.0024251668],"about_ca_topic_score_codex":0.0015334507,"about_ca_topic_score_gemma":0.0031144666,"teacher_disagreement_score":0.049793772,"about_ca_system_score_codex":0.0011560633,"about_ca_system_score_gemma":0.0013086294,"threshold_uncertainty_score":0.26333773},"labels":[],"label_agreement":null},{"id":"W3176893767","doi":"10.17504/protocols.io.buhwnt7e","title":"SARS-CoV-2 ENA submission workflow + guidance for structuring and releasing metadata v1","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Metadata; Workflow; Computer science; Protocol (science); World Wide Web; Biology; Database","score_opus":0.048604500684308875,"score_gpt":0.32415883663765094,"score_spread":0.27555433595334206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3176893767","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00375706,0.00079198164,0.35114247,0.0056941737,0.003146225,0.014257438,0.30930415,0.22954059,0.08236589],"genre_scores_gemma":[0.009340126,0.0012651467,0.37469804,0.004653616,0.00083566684,0.020293977,0.41712123,0.095717296,0.07607493],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9849555,0.0040780394,0.0034656131,0.0016797973,0.0046870476,0.0011340276],"domain_scores_gemma":[0.95948154,0.011009319,0.0023524924,0.013546544,0.011753402,0.0018567236],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.0313922,0.0022500106,0.0016932269,0.005783544,0.00327712,0.010606877,0.00354784,0.003725358,0.19441813],"category_scores_gemma":[0.06443402,0.0028660425,0.0022276763,0.003895298,0.0012284693,0.0065422426,0.008943452,0.004366895,0.366276],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062068406,0.00014891104,0.0013866336,0.0013889021,0.000046598296,0.00030870835,0.00064202154,0.00072269246,0.0063464586,0.011094661,0.8931406,0.08415307],"study_design_scores_gemma":[0.00010147485,0.000056102024,0.0009873548,0.0007978517,0.000014543233,0.00022055407,0.0002446457,0.00086285657,0.007159451,0.0066344594,0.9828183,0.00010242074],"about_ca_topic_score_codex":0.0061024358,"about_ca_topic_score_gemma":0.004873558,"teacher_disagreement_score":0.99645215,"about_ca_system_score_codex":0.0032933177,"about_ca_system_score_gemma":0.010191589,"threshold_uncertainty_score":0.6503935},"labels":[],"label_agreement":null},{"id":"W3181217566","doi":"10.1158/1538-7445.am2021-449","title":"Abstract 449: A standard operating procedure for the curation of gene fusions","year":2021,"lang":"en","type":"article","venue":"Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network","funders":"","keywords":"Fusion gene; Interpretation (philosophy); Computational biology; Interoperability; Set (abstract data type); Gene; Biology; Computer science; Bioinformatics; Data science; Genetics; World Wide Web","score_opus":0.10215346907557346,"score_gpt":0.44957096225736637,"score_spread":0.34741749318179294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3181217566","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038537297,0.00019874405,0.8216925,0.00081689295,0.0005044861,0.0035602057,0.008677279,0.15495673,0.0057394016],"genre_scores_gemma":[0.024855586,0.00026992918,0.90220994,0.0011968061,0.0002819891,0.008029174,0.022330912,0.03401888,0.0068067666],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.93449736,0.020264098,0.014992907,0.008369173,0.019634802,0.0022416674],"domain_scores_gemma":[0.8674543,0.062339675,0.009510911,0.02428079,0.033469286,0.0029450296],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.057875197,0.003384084,0.0022134979,0.010298322,0.002778818,0.007411684,0.004277869,0.003678462,0.044950742],"category_scores_gemma":[0.13786878,0.0021419413,0.0041064643,0.005279766,0.0032539107,0.004663086,0.009256734,0.0056327777,0.05345745],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005388331,0.00057437905,0.013660853,0.00464664,0.0008211891,0.0052329814,0.0054881466,0.0070211734,0.099416,0.039004736,0.3136981,0.5050475],"study_design_scores_gemma":[0.00061848963,0.00077891344,0.009990078,0.0014890747,0.0002545071,0.0030477608,0.0012394432,0.053966638,0.17064774,0.04637207,0.7106872,0.0009080269],"about_ca_topic_score_codex":0.0030477145,"about_ca_topic_score_gemma":0.0019753352,"teacher_disagreement_score":0.057875197,"about_ca_system_score_codex":0.003009226,"about_ca_system_score_gemma":0.011984964,"threshold_uncertainty_score":0.306077},"labels":[],"label_agreement":null},{"id":"W3181498676","doi":"10.1093/database/baab040","title":"Standardization of assay representation in the Ontology for Biomedical Investigations","year":2021,"lang":"en","type":"review","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Institute of Allergy and Infectious Diseases; National Human Genome Research Institute; Bill and Melinda Gates Foundation; National Institutes of Health; National Science Foundation","keywords":"Standardization; Ontology; Computer science; Process (computing); Open Biomedical Ontologies; Software engineering; Representation (politics); Hierarchy; Process ontology; Information retrieval; Ontology Inference Layer; OWL-S; Term (time); World Wide Web; Semantic Web; Suggested Upper Merged Ontology; Programming language; Semantic Web Stack","score_opus":0.11938547991577698,"score_gpt":0.4300302172784989,"score_spread":0.3106447373627219,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3181498676","genre_codex":"review","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020294618,0.8701327,0.093468636,0.00958527,0.003373534,0.00039609344,0.0011276198,0.00039220433,0.019494465],"genre_scores_gemma":[0.012446586,0.8472579,0.12530643,0.005439796,0.0007192548,0.0005598336,0.0036813535,0.0001550638,0.0044338047],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99654007,0.00095385185,0.0007682963,0.0004586613,0.0011306343,0.00014860554],"domain_scores_gemma":[0.9940519,0.0027977908,0.00052856305,0.00073272386,0.0017429687,0.00014603714],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014602051,0.0009517474,0.0021333606,0.0066903858,0.0006307993,0.003886955,0.0030668438,0.0021519996,0.0016261349],"category_scores_gemma":[0.009668819,0.00057605415,0.0016838834,0.007293423,0.0023020424,0.0057595912,0.0025907836,0.0049639544,0.0018518885],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006271723,0.00011871926,0.00053469726,0.018993568,0.00019504182,0.0003148557,0.0006291843,0.0011044288,0.0057113348,0.16941312,0.021656342,0.781266],"study_design_scores_gemma":[0.000009901303,0.000031054693,0.0005125029,0.00510426,0.00014933059,0.000609671,0.00018623222,0.0003335836,0.0022560991,0.0145310275,0.97624105,0.000035346107],"about_ca_topic_score_codex":0.0048957113,"about_ca_topic_score_gemma":0.0035263563,"teacher_disagreement_score":0.014602051,"about_ca_system_score_codex":0.0039474647,"about_ca_system_score_gemma":0.0101489145,"threshold_uncertainty_score":0.07722396},"labels":[],"label_agreement":null},{"id":"W3182655775","doi":"10.2196/27970","title":"A Natural Language Processing–Assisted Extraction System for Gleason Scores: Development and Usability Study","year":2021,"lang":"en","type":"article","venue":"JMIR Cancer","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Cancer Institute; Abramson Cancer Center; University of Pennsylvania Health System; University of Pennsylvania","keywords":"Artificial intelligence; Medicine; Prostate cancer; Cohort; Natural language processing; Categorization; Computer science; Cancer; Pathology; Internal medicine","score_opus":0.022899252833520907,"score_gpt":0.35244826165148824,"score_spread":0.32954900881796734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3182655775","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85376245,0.0005093089,0.1252245,0.00062915037,0.00011895657,0.0063027907,0.0018658245,0.009354264,0.002232699],"genre_scores_gemma":[0.63703454,0.00044496285,0.35314107,0.00047605523,0.00006724921,0.0030771457,0.0035847765,0.0007224675,0.001451829],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99080616,0.005597915,0.001098021,0.00095899595,0.001314126,0.00022479426],"domain_scores_gemma":[0.9361591,0.0503925,0.00107142,0.0029253503,0.008802996,0.0006485861],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025346909,0.00088145875,0.0005875724,0.0013049745,0.00049247895,0.0013537813,0.0014288607,0.0008067295,0.0019934657],"category_scores_gemma":[0.04875221,0.00052299595,0.00076111837,0.0007415398,0.00054159976,0.0027717815,0.0014018987,0.00082541996,0.00095077883],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035648753,0.005603826,0.10662736,0.004394179,0.00083636906,0.0027499567,0.016176248,0.007023971,0.07221696,0.0010943158,0.015137534,0.76457447],"study_design_scores_gemma":[0.0040376307,0.020117085,0.30264875,0.0017106632,0.0019381028,0.007967292,0.010665067,0.3911724,0.1781602,0.0031425017,0.07740656,0.0010337261],"about_ca_topic_score_codex":0.001993196,"about_ca_topic_score_gemma":0.002323944,"teacher_disagreement_score":0.025346909,"about_ca_system_score_codex":0.0006640681,"about_ca_system_score_gemma":0.0013502673,"threshold_uncertainty_score":0.13404882},"labels":[],"label_agreement":null},{"id":"W3182801416","doi":"10.1007/s40037-021-00675-8","title":"When names are on the line: Negotiating authorship with your team","year":2021,"lang":"en","type":"editorial","venue":"Perspectives on Medical Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Negotiation; Medical education; Line (geometry); Computer science; Psychology; Medicine; Sociology; Mathematics","score_opus":0.01620751826470444,"score_gpt":0.3240968551179139,"score_spread":0.30788933685320946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3182801416","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000031958338,0.001268809,0.0003920845,0.1343311,0.86281115,0.000013499001,0.00002757536,0.000048338465,0.0010754875],"genre_scores_gemma":[0.0008063128,0.0014575147,0.00045360075,0.06728256,0.9193076,0.000043225937,0.000028668524,0.00009287912,0.010527641],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96542037,0.010223345,0.004581481,0.0028951836,0.014909518,0.001970076],"domain_scores_gemma":[0.80282086,0.11309137,0.008118381,0.005366646,0.0479029,0.022699865],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.0469952,0.0022955965,0.004061374,0.0034787285,0.0070295506,0.02088859,0.004316823,0.030269418,0.017768476],"category_scores_gemma":[0.17730305,0.0015315987,0.0026406404,0.0021934302,0.00620804,0.011856511,0.004442452,0.03885017,0.013728349],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027018206,0.000008932468,0.000028787048,0.000083000166,0.00001800101,0.000094191084,0.000047623183,0.000017377632,0.000028023978,0.000887038,0.99472165,0.0040383367],"study_design_scores_gemma":[0.00006152556,0.00001604646,0.0001367072,0.00047255895,0.00005073854,0.00018732315,0.00020497163,0.00022778264,0.000103887454,0.0047254274,0.9937736,0.000039409384],"about_ca_topic_score_codex":0.0021872523,"about_ca_topic_score_gemma":0.00833974,"teacher_disagreement_score":0.96973056,"about_ca_system_score_codex":0.005878301,"about_ca_system_score_gemma":0.010563215,"threshold_uncertainty_score":0.24853736},"labels":[],"label_agreement":null},{"id":"W3182985492","doi":"10.3390/ijerph18147355","title":"A Health eLearning Ontology and Procedural Reasoning Approach for Developing Personalized Courses to Teach Patients about Their Medical Condition and Treatment","year":2021,"lang":"en","type":"article","venue":"International Journal of Environmental Research and Public Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; University of Ottawa","funders":"Politechnika Poznańska","keywords":"Operationalization; Computer science; Ontology; Comprehension; Learning styles; Personalized medicine; Medical education; Multimedia; Medicine; Psychology; Mathematics education; Bioinformatics","score_opus":0.05694594985870276,"score_gpt":0.3981072378833512,"score_spread":0.3411612880246484,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3182985492","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009885113,0.00010466515,0.97776437,0.0014108185,0.00003092409,0.00048924476,0.00014834267,0.0006918845,0.009474683],"genre_scores_gemma":[0.057563808,0.00010469357,0.9389597,0.00017807193,0.000009310356,0.0004736351,0.00026072882,0.000059143968,0.00239094],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964575,0.0017923126,0.00039293297,0.0006300219,0.0005456146,0.00018154114],"domain_scores_gemma":[0.99434066,0.0032898842,0.000483619,0.0009601384,0.00057555834,0.00035001934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006744987,0.0006836701,0.00034619455,0.002113716,0.00115598,0.0031736286,0.002486874,0.0014061253,0.00443871],"category_scores_gemma":[0.008999583,0.00049164525,0.001366658,0.0013139654,0.0030543155,0.0046387273,0.0035518797,0.0017469707,0.0008012101],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011121883,0.001383297,0.008233202,0.00086597377,0.000118101,0.0005032591,0.011935127,0.027298378,0.009819213,0.58071935,0.005554619,0.35345832],"study_design_scores_gemma":[0.00022582553,0.0005853738,0.005160994,0.000866058,0.0003018347,0.0013704373,0.004859084,0.18974939,0.021372506,0.49779373,0.27755082,0.00016398726],"about_ca_topic_score_codex":0.0028601668,"about_ca_topic_score_gemma":0.004335728,"teacher_disagreement_score":0.006744987,"about_ca_system_score_codex":0.0023357202,"about_ca_system_score_gemma":0.0048212195,"threshold_uncertainty_score":0.035671294},"labels":[],"label_agreement":null},{"id":"W3185277452","doi":"10.1016/j.ijmedinf.2021.104539","title":"Comparison of terminology mapping methods for nursing wound care knowledge representation","year":2021,"lang":"en","type":"article","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vancouver Coastal Health; University of British Columbia","funders":"","keywords":"SNOMED CT; Terminology; Concordance; Systematized Nomenclature of Medicine; Representativeness heuristic; Computer science; Context (archaeology); Wound care; Controlled vocabulary; Artificial intelligence; Information retrieval; Natural language processing; Data mining; Medicine; Mathematics; Statistics; Linguistics; Surgery","score_opus":0.09967698394488882,"score_gpt":0.5195603411493498,"score_spread":0.41988335720446096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3185277452","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19487624,0.011296533,0.75446635,0.0021985595,0.0006303203,0.0015942261,0.010860387,0.0069006756,0.017176775],"genre_scores_gemma":[0.3640397,0.0055323173,0.6037046,0.0003287522,0.0000997997,0.0009237902,0.020908445,0.0005198335,0.003942862],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9946097,0.0022421943,0.0009898891,0.0005353767,0.0014191301,0.00020365289],"domain_scores_gemma":[0.981006,0.013719974,0.00055189634,0.00130472,0.0031740984,0.00024330613],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007416632,0.0009113754,0.0010377375,0.010450312,0.0009928955,0.0044207405,0.0017389893,0.0010891431,0.0031889903],"category_scores_gemma":[0.02408776,0.0003165588,0.0020214394,0.008847627,0.00045015872,0.0054133814,0.0021921801,0.0011101306,0.0009722892],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010359931,0.0006478129,0.009785606,0.00190835,0.0007816094,0.00018079045,0.0016307944,0.01519267,0.005074985,0.016135976,0.009603669,0.9380218],"study_design_scores_gemma":[0.0006188482,0.0015278236,0.052214157,0.0024335294,0.0029707376,0.0017924557,0.009860477,0.6934412,0.031079054,0.10178342,0.101881795,0.00039645928],"about_ca_topic_score_codex":0.009010158,"about_ca_topic_score_gemma":0.008295151,"teacher_disagreement_score":0.010450312,"about_ca_system_score_codex":0.0014489919,"about_ca_system_score_gemma":0.0032490538,"threshold_uncertainty_score":0.039223373},"labels":[],"label_agreement":null},{"id":"W3187314870","doi":"10.3390/info12080317","title":"Design of Generalized Search Interfaces for Health Informatics","year":2021,"lang":"en","type":"article","venue":"Information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Ontology; Health informatics; Vocabulary; Interface (matter); Workflow; Domain (mathematical analysis); Controlled vocabulary; Information retrieval; Informatics; User interface; Set (abstract data type); Plug-in; Human–computer interaction; Data science; World Wide Web; Database; Public health; Engineering; Medicine; Programming language","score_opus":0.04229474773318576,"score_gpt":0.3360776097010882,"score_spread":0.29378286196790243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3187314870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014608789,0.00017325373,0.9743593,0.0003231544,0.000035149256,0.0006542211,0.00020843714,0.0071711903,0.0024664125],"genre_scores_gemma":[0.11167395,0.00017095277,0.882213,0.00034434962,0.000022108452,0.0010411199,0.0007206171,0.0010015947,0.0028122894],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99464864,0.0024281223,0.0008138459,0.00078516424,0.0010776022,0.00024662243],"domain_scores_gemma":[0.9882719,0.0076091685,0.0005679146,0.0016946264,0.0014965505,0.00035975935],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007063594,0.0011012182,0.0009077408,0.0017139316,0.00078100205,0.003992536,0.0024421273,0.0016046368,0.005992497],"category_scores_gemma":[0.023070848,0.00083154725,0.0011568106,0.0011330104,0.0014419119,0.0053543444,0.0036294835,0.0011487155,0.0016153221],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020233702,0.00059914123,0.0076395907,0.0037469517,0.00042828333,0.0021384843,0.017999465,0.042146493,0.07201203,0.28798914,0.025802063,0.537475],"study_design_scores_gemma":[0.0006059423,0.00066263555,0.0030610939,0.000671211,0.0003575394,0.0016090531,0.0035296015,0.5749235,0.04827332,0.20269158,0.16336472,0.00024978942],"about_ca_topic_score_codex":0.0020221544,"about_ca_topic_score_gemma":0.002242103,"teacher_disagreement_score":0.007063594,"about_ca_system_score_codex":0.0009830283,"about_ca_system_score_gemma":0.001789637,"threshold_uncertainty_score":0.037356317},"labels":[],"label_agreement":null},{"id":"W3189155455","doi":"10.2139/ssrn.3199336","title":"Semantic Web Infrastructure for Fungal Enzyme Biotechnologists","year":2006,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; World Wide Web; Medicine","score_opus":0.004451410023825407,"score_gpt":0.23285911530816902,"score_spread":0.22840770528434362,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3189155455","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028333584,0.001035878,0.80367655,0.0026415836,0.00030173684,0.00074726966,0.04166659,0.08902117,0.032575626],"genre_scores_gemma":[0.17024235,0.002368064,0.55857277,0.0011362648,0.00017473026,0.00094352494,0.24622957,0.00559487,0.014737884],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9992508,0.00012295459,0.00015243003,0.00012644933,0.00026386773,0.00008341728],"domain_scores_gemma":[0.9981812,0.00042030722,0.00022105942,0.0005923716,0.00038386157,0.00020118288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024888038,0.00073390716,0.00073221343,0.0043117474,0.0014671583,0.002773022,0.0013662701,0.0014805257,0.0058880397],"category_scores_gemma":[0.0034030427,0.00055699516,0.0012761621,0.004303217,0.00059049006,0.00613893,0.0028381178,0.0013487616,0.0046643787],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00169862,0.0012401986,0.00896608,0.0027808114,0.0003535619,0.0019848864,0.0018697414,0.011815746,0.04531762,0.26696238,0.20671496,0.45029545],"study_design_scores_gemma":[0.00021709086,0.00012836841,0.0066895992,0.00079799845,0.0003917968,0.001097733,0.0008199771,0.09599284,0.040827062,0.19147868,0.66142404,0.0001349167],"about_ca_topic_score_codex":0.0047481866,"about_ca_topic_score_gemma":0.00462954,"teacher_disagreement_score":0.0058880397,"about_ca_system_score_codex":0.0011799538,"about_ca_system_score_gemma":0.0035391538,"threshold_uncertainty_score":0.019697487},"labels":[],"label_agreement":null},{"id":"W3191528200","doi":"10.1109/bhi50953.2021.9508585","title":"Chatsum: An Intelligent Medical Chat Summarization Tool","year":2021,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Automatic summarization; Computer science; Upload; World Wide Web; Service (business); Usability; Multimedia; The Internet; Human–computer interaction; Information retrieval","score_opus":0.0184921892785718,"score_gpt":0.2957641800822204,"score_spread":0.27727199080364856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3191528200","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.032432973,0.0012423656,0.6305557,0.0011556019,0.00057805696,0.0014825225,0.030217467,0.29713106,0.00520423],"genre_scores_gemma":[0.1466652,0.00071938505,0.7769182,0.0005536232,0.00048505032,0.0015419351,0.05630756,0.0051453114,0.011663722],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988111,0.00039058662,0.00020509853,0.00026264615,0.0002840787,0.00004646072],"domain_scores_gemma":[0.9936621,0.0039715883,0.0006578697,0.00036365885,0.0010912415,0.00025355953],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019439621,0.0018315337,0.0007129551,0.004237565,0.0006256245,0.001161624,0.0012942336,0.00089383096,0.01135932],"category_scores_gemma":[0.010245526,0.00043361806,0.00077458046,0.0013199496,0.0002232124,0.0015015668,0.0011664778,0.0008654853,0.0041912957],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001603621,0.00037010745,0.0043338975,0.0038652427,0.00039981958,0.0010299472,0.0025164906,0.0061268546,0.04615968,0.0029367579,0.17289466,0.75776297],"study_design_scores_gemma":[0.00089545996,0.0016803419,0.025445161,0.00095524563,0.00083159516,0.0030123768,0.0026724357,0.42622915,0.1380438,0.01703754,0.3826213,0.00057563646],"about_ca_topic_score_codex":0.0015565465,"about_ca_topic_score_gemma":0.0029590183,"teacher_disagreement_score":0.01135932,"about_ca_system_score_codex":0.00050340133,"about_ca_system_score_gemma":0.000797469,"threshold_uncertainty_score":0.038000762},"labels":[],"label_agreement":null},{"id":"W3194665423","doi":"10.2196/28212","title":"Matching Biomedical Ontologies: Construction of Matching Clues and Systematic Evaluation of Different Combinations of Matchers","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Ontology alignment; Matching (statistics); Ontology; Terminology; Interoperability; Data mining; Machine learning; Artificial intelligence; Semantic Web; Ontology-based data integration; Mathematics","score_opus":0.02283905174950861,"score_gpt":0.32163936823827316,"score_spread":0.29880031648876454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3194665423","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26402614,0.0036942388,0.7188286,0.0007418944,0.00020419675,0.001279967,0.0010952889,0.004516247,0.0056133955],"genre_scores_gemma":[0.36024746,0.00087957206,0.6353618,0.00014283754,0.000036932066,0.0003249097,0.0018387587,0.00024137617,0.0009263684],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9847852,0.0050679096,0.0018654879,0.0027014446,0.0051095015,0.0004704327],"domain_scores_gemma":[0.97550166,0.014421509,0.0019226798,0.0032300397,0.004304478,0.0006196347],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011042351,0.0015776865,0.0014701036,0.007848497,0.0013423132,0.002361237,0.0022113286,0.0017133126,0.0025204937],"category_scores_gemma":[0.053102486,0.0005208699,0.0014276807,0.0042120484,0.0012276033,0.006326134,0.004168831,0.0012358463,0.0008674043],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015912864,0.0005851604,0.023135424,0.002407915,0.00052445114,0.00045505457,0.0012773348,0.039797142,0.032583423,0.01928123,0.0060120164,0.8723495],"study_design_scores_gemma":[0.00029667196,0.001938208,0.013651968,0.0006042984,0.0011977575,0.0018248353,0.0034977144,0.79826266,0.112146124,0.038645778,0.027704164,0.000229766],"about_ca_topic_score_codex":0.0028584825,"about_ca_topic_score_gemma":0.003171953,"teacher_disagreement_score":0.011042351,"about_ca_system_score_codex":0.0014575066,"about_ca_system_score_gemma":0.003277464,"threshold_uncertainty_score":0.058398247},"labels":[],"label_agreement":null},{"id":"W3196092690","doi":"10.1002/wfs2.1441","title":"Towards another paradigm for forensic science?","year":2021,"lang":"en","type":"article","venue":"Wiley Interdisciplinary Reviews Forensic Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Trois-Rivières","funders":"","keywords":"Transparency (behavior); Epistemology; Abductive reasoning; Computer science; Semiotics; Data science; Artificial intelligence; Interpretation (philosophy); Psychology; Cognitive science; Philosophy; Computer security","score_opus":0.04147954242890188,"score_gpt":0.35149990485541144,"score_spread":0.31002036242650954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3196092690","genre_codex":"commentary","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005348203,0.044126034,0.14601299,0.6625678,0.0069192676,0.00008292252,0.000091900896,0.00018664756,0.13466421],"genre_scores_gemma":[0.65892327,0.044991717,0.1517844,0.092228934,0.015815642,0.00089382456,0.00020256736,0.00037445745,0.034785207],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98243064,0.012716845,0.0004888631,0.0013395839,0.002505636,0.0005184343],"domain_scores_gemma":[0.98280084,0.01061887,0.00088977884,0.0019448777,0.0027462924,0.0009992933],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.029749665,0.00096434227,0.0015605511,0.00572985,0.007468415,0.017352663,0.0032546562,0.008909413,0.0048671286],"category_scores_gemma":[0.014834095,0.00053680723,0.001126384,0.002495562,0.09397452,0.02901273,0.00765132,0.013166217,0.0015859584],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000022644356,0.0000037081722,0.000012667003,0.000022698267,0.0000014388743,0.0000114770955,0.00053059176,0.000048984562,0.000011578346,0.9963685,0.0014732534,0.0015128291],"study_design_scores_gemma":[0.0000035282092,0.0000027710482,0.000011671868,0.000079080615,8.064795e-7,0.000021200169,0.00047721493,0.00016688714,0.000015589372,0.97939175,0.019825555,0.0000038526164],"about_ca_topic_score_codex":0.0034512992,"about_ca_topic_score_gemma":0.001941112,"teacher_disagreement_score":0.9925316,"about_ca_system_score_codex":0.012286456,"about_ca_system_score_gemma":0.012027612,"threshold_uncertainty_score":0.15733314},"labels":[],"label_agreement":null},{"id":"W3197008538","doi":"10.1038/s41467-021-25578-4","title":"Automatically disambiguating medical acronyms with ontology-aware deep learning","year":2021,"lang":"en","type":"article","venue":"Nature Communications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; Hospital for Sick Children; Toronto Rehabilitation Institute; University of Toronto; University Health Network","funders":"Hospital for Sick Children","keywords":"Computer science; Workflow; Task (project management); Context (archaeology); Artificial intelligence; Ontology; Training set; Set (abstract data type); Test set; Labeled data; Natural language processing; Machine learning; Information retrieval; Database","score_opus":0.01294677349322768,"score_gpt":0.3175330017944535,"score_spread":0.3045862283012258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3197008538","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16542752,0.008323068,0.7392828,0.004650509,0.00196154,0.00049665454,0.02380896,0.039094992,0.01695397],"genre_scores_gemma":[0.42187834,0.0019290729,0.5191171,0.0017154863,0.00044780484,0.0002672779,0.045215108,0.00066987955,0.00875987],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99905163,0.00018106244,0.00011057521,0.00041922994,0.00016352153,0.000074081305],"domain_scores_gemma":[0.99838376,0.0007172533,0.00022430535,0.00029038565,0.00029632094,0.000087888504],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011280349,0.0014004844,0.0006682602,0.0039236546,0.0006805343,0.0013277619,0.0015341779,0.0012650171,0.003036016],"category_scores_gemma":[0.003974604,0.00039692217,0.0012165989,0.0025277296,0.0006976699,0.0024116016,0.0021131127,0.0019457836,0.0028176552],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046206889,0.00034318044,0.011560584,0.00075801084,0.00019135234,0.00069687265,0.00044272208,0.03971218,0.019894907,0.009783482,0.0954554,0.82069916],"study_design_scores_gemma":[0.00015436711,0.0001483556,0.00624016,0.0003020645,0.00021411788,0.00087666255,0.0005810678,0.81802654,0.02769422,0.04653297,0.09912572,0.000103805636],"about_ca_topic_score_codex":0.008074093,"about_ca_topic_score_gemma":0.017640082,"teacher_disagreement_score":0.008074093,"about_ca_system_score_codex":0.0013769271,"about_ca_system_score_gemma":0.002414478,"threshold_uncertainty_score":0.016054213},"labels":[],"label_agreement":null},{"id":"W3199267631","doi":"10.2196/29398","title":"A Deep Learning Approach to Refine the Identification of High-Quality Clinical Research Articles From the Biomedical Literature: Protocol for Algorithm Development and Validation","year":2021,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"McMaster University; Impact","funders":"","keywords":"Computer science; Machine learning; Hyperparameter; Artificial intelligence; Relevance (law); Identification (biology); Deep learning; Protocol (science); Data mining; Algorithm; Information retrieval; Medicine","score_opus":0.3789286136609896,"score_gpt":0.5897857567159527,"score_spread":0.21085714305496306,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3199267631","genre_codex":"methods","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009683521,0.0013028352,0.6563973,0.0012577472,0.00047393696,0.30023158,0.017291779,0.0071309325,0.006230304],"genre_scores_gemma":[0.009791954,0.00034947347,0.580285,0.00039777387,0.00004085338,0.4024033,0.004810309,0.00037152934,0.0015498219],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9797934,0.009934306,0.0036501058,0.002474754,0.0035743196,0.0005731297],"domain_scores_gemma":[0.92213887,0.037702132,0.0042396844,0.015514179,0.019189177,0.0012159608],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.061637737,0.0036157605,0.001924919,0.004372644,0.002182787,0.003031754,0.0048907683,0.0037930966,0.039989226],"category_scores_gemma":[0.13182223,0.002461275,0.0037451463,0.0026998431,0.002933876,0.0023778297,0.004156449,0.0062154285,0.011880726],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.011681034,0.0045675607,0.007387927,0.021952577,0.0019131186,0.0011739368,0.0014763774,0.06810937,0.022808593,0.0320312,0.10540814,0.7214901],"study_design_scores_gemma":[0.029628862,0.009092642,0.01045769,0.01432345,0.0024154286,0.0025526066,0.0009941993,0.31860372,0.09707118,0.10582521,0.4080134,0.001021589],"about_ca_topic_score_codex":0.0032924053,"about_ca_topic_score_gemma":0.0052225036,"teacher_disagreement_score":0.93836224,"about_ca_system_score_codex":0.003904037,"about_ca_system_score_gemma":0.02302173,"threshold_uncertainty_score":0.32597542},"labels":[],"label_agreement":null},{"id":"W3201217389","doi":"10.2196/27550","title":"Toward Data-Driven Radiation Oncology Using Standardized Terminology as a Starting Point: Cross-sectional Study","year":2021,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Terminology; Unified Medical Language System; Interoperability; Descriptive statistics; Computer science; Radiation oncology; Medicine; Medical physics; Radiation therapy; Information retrieval; Internal medicine; World Wide Web; Statistics; Linguistics","score_opus":0.18590004366692228,"score_gpt":0.5105762371406873,"score_spread":0.3246761934737651,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3201217389","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9939319,0.0013571897,0.0017294319,0.0004207135,0.000022947164,0.00015369552,0.0006649123,0.00000774016,0.0017114535],"genre_scores_gemma":[0.9970913,0.00048419452,0.0012206956,0.00020689073,0.0000120958575,0.00024652557,0.0005149596,0.000016498267,0.00020687908],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9878172,0.00554413,0.0018593805,0.001759518,0.002314652,0.00070513715],"domain_scores_gemma":[0.9480035,0.02243116,0.016620183,0.0029182837,0.008683776,0.0013431493],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027252622,0.00029460245,0.0005699258,0.0041719577,0.0012533881,0.0032902074,0.0008969803,0.0008644093,0.0024989198],"category_scores_gemma":[0.0493476,0.0007047998,0.00096298975,0.0057889833,0.0014076333,0.0047022738,0.003015821,0.0016560993,0.000512959],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004095756,0.000071469745,0.9821584,0.00018635012,0.00007681293,0.00007953199,0.011932553,0.000042206382,0.00013531311,0.0007814908,0.0004522933,0.004042507],"study_design_scores_gemma":[0.0000110337705,0.00035698878,0.9256368,0.0007448705,0.00016912787,0.0008880414,0.06222582,0.000707709,0.0003389555,0.00067412277,0.0082074115,0.000039052367],"about_ca_topic_score_codex":0.0070445403,"about_ca_topic_score_gemma":0.006346534,"teacher_disagreement_score":0.027252622,"about_ca_system_score_codex":0.0017457238,"about_ca_system_score_gemma":0.0029888563,"threshold_uncertainty_score":0.14412743},"labels":[],"label_agreement":null},{"id":"W3203751410","doi":"10.1093/database/baab062","title":"Classifying domain-specific text documents containing ambiguous keywords","year":2021,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development","keywords":"Computer science; Set (abstract data type); Domain (mathematical analysis); Information retrieval; Artificial intelligence; Naive Bayes classifier; Overfitting; Identifier; Machine learning; Support vector machine; Ambiguity; Data mining; Artificial neural network","score_opus":0.027044701699820522,"score_gpt":0.290375614861075,"score_spread":0.26333091316125445,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3203751410","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41285816,0.11671267,0.15473124,0.01113487,0.0046367166,0.0067344606,0.23573327,0.013085415,0.04437329],"genre_scores_gemma":[0.30902308,0.032948438,0.51168305,0.0028687215,0.0014586045,0.002540709,0.12884879,0.0008441078,0.009784491],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99483657,0.0010444738,0.0017965898,0.00079994317,0.0013637651,0.00015870672],"domain_scores_gemma":[0.93145716,0.050163794,0.0066127367,0.0028036013,0.0082093,0.0007533582],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00884917,0.0010974831,0.0017116887,0.035691865,0.0012953893,0.0045972858,0.0011123196,0.0014949875,0.0067792516],"category_scores_gemma":[0.03963737,0.00031444276,0.0016046237,0.02888635,0.00055102084,0.0028796485,0.0011377129,0.0005608382,0.003566945],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015703838,0.0004481158,0.037727963,0.05686154,0.0012508805,0.00480184,0.0020073126,0.0028222543,0.050806627,0.007982881,0.08699573,0.74672455],"study_design_scores_gemma":[0.00083298434,0.0015520273,0.10284069,0.013410776,0.004321645,0.013380378,0.009624744,0.022519663,0.11479493,0.038255874,0.6779065,0.0005597386],"about_ca_topic_score_codex":0.0017132073,"about_ca_topic_score_gemma":0.0036635222,"teacher_disagreement_score":0.035691865,"about_ca_system_score_codex":0.0011187114,"about_ca_system_score_gemma":0.0036396284,"threshold_uncertainty_score":0.04679948},"labels":[],"label_agreement":null},{"id":"W3205474595","doi":"10.1504/ijiids.2022.10041973","title":"Supporting user-centred ontology visualisation: predictive analytics using eye gaze to enhance human-ontology interaction","year":2021,"lang":"en","type":"article","venue":"International Journal of Intelligent Information and Database Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of British Columbia","funders":"","keywords":"Computer science; Ontology; Human–computer interaction; Visual analytics; Gaze; Visualization; Eye tracking; Analytics; Data science; Artificial intelligence","score_opus":0.0399893152729184,"score_gpt":0.39771629311566736,"score_spread":0.357726977842749,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3205474595","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.66815764,0.0010653149,0.31324884,0.00055146107,0.00009677999,0.00035108987,0.0019205371,0.006148761,0.008459572],"genre_scores_gemma":[0.9339579,0.00032569954,0.06340065,0.00005103639,0.000029931696,0.00011552179,0.0005097045,0.00018624162,0.0014233205],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999411,0.00023113206,0.000023519015,0.000116478484,0.00016136865,0.00005658061],"domain_scores_gemma":[0.99572766,0.0030712727,0.0003506375,0.00030286165,0.0004148718,0.00013273308],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010403864,0.00068724016,0.00046551073,0.0019165657,0.00033925133,0.0015627942,0.00047803952,0.00060932414,0.003277944],"category_scores_gemma":[0.009879142,0.00022057608,0.00031210182,0.00086019345,0.00026154864,0.0013731021,0.0014592501,0.00058859866,0.0009375125],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027259863,0.00055257214,0.08659836,0.0017200657,0.00026529699,0.0005825365,0.014494286,0.020626146,0.20110317,0.0033129044,0.009204707,0.658814],"study_design_scores_gemma":[0.00014742375,0.0017638602,0.28961036,0.0005343311,0.00028137458,0.0014793875,0.006871378,0.5732742,0.08757464,0.017282907,0.020704174,0.00047597417],"about_ca_topic_score_codex":0.0044483924,"about_ca_topic_score_gemma":0.006214368,"teacher_disagreement_score":0.0044483924,"about_ca_system_score_codex":0.00036444725,"about_ca_system_score_gemma":0.00042447352,"threshold_uncertainty_score":0.010965884},"labels":[],"label_agreement":null},{"id":"W3206073083","doi":"10.1503/cmaj.210877-f","title":"Optimiser les données accessibles dans le portail de renseignements cliniques de Santé Canada","year":2021,"lang":"fr","type":"article","venue":"Canadian Medical Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"National Center for Advancing Translational Sciences; Agency for Healthcare Research and Quality; National Heart, Lung, and Blood Institute; U.S. Food and Drug Administration","keywords":"Political science; Medicine","score_opus":0.012836778519520023,"score_gpt":0.2645669368031975,"score_spread":0.25173015828367745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3206073083","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38830474,0.024442902,0.30103812,0.063492276,0.001116951,0.0029319134,0.1557398,0.013697777,0.049235415],"genre_scores_gemma":[0.54742503,0.0074785594,0.32377213,0.0023276545,0.00019722836,0.000577974,0.104841754,0.0010897351,0.012289787],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97295195,0.006233694,0.001929262,0.0029735817,0.01454075,0.0013706726],"domain_scores_gemma":[0.94411284,0.023034943,0.0023315602,0.003638834,0.025313403,0.0015684079],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.021585401,0.0013842185,0.0013820878,0.008164496,0.003395139,0.011706248,0.0023347049,0.0017108357,0.0028583845],"category_scores_gemma":[0.076257035,0.0007675012,0.0018748473,0.009661621,0.0017194804,0.005558909,0.0037443887,0.0033535045,0.0011704832],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003197134,0.00091319217,0.13912728,0.0065005086,0.0021349133,0.002751048,0.01107437,0.09585714,0.033418015,0.044020526,0.122253984,0.5387519],"study_design_scores_gemma":[0.0004475413,0.0004941937,0.1259693,0.0040768785,0.0018131675,0.0012319987,0.019562304,0.40336865,0.048509978,0.053420655,0.34045333,0.00065192336],"about_ca_topic_score_codex":0.7106581,"about_ca_topic_score_gemma":0.7518861,"teacher_disagreement_score":0.9976653,"about_ca_system_score_codex":0.015436042,"about_ca_system_score_gemma":0.05472117,"threshold_uncertainty_score":0.5820918},"labels":[],"label_agreement":null},{"id":"W3207456493","doi":"10.1101/2021.10.13.464308","title":"PSEA: A phenotypic similarity ensemble approach for prioritizes candidate genes to aid mendelian disease diagnosis","year":2021,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wycliffe College","funders":"","keywords":"Phenotype; Computational biology; Similarity (geometry); Genetics; Biology; Gene; Computer science; Artificial intelligence","score_opus":0.020482848119446898,"score_gpt":0.25225134332040877,"score_spread":0.23176849520096188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207456493","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10523111,0.0018197726,0.86703974,0.0008000446,0.0001735137,0.00033483715,0.010184122,0.010266905,0.0041499455],"genre_scores_gemma":[0.51177216,0.00074161805,0.46240523,0.00046041654,0.0002711292,0.0004888626,0.019766571,0.0009431174,0.0031508561],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986211,0.00033010746,0.000094184485,0.00045781498,0.00041563695,0.00008119149],"domain_scores_gemma":[0.99846494,0.0007766543,0.00016061586,0.00017954703,0.00030367335,0.0001146194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002059692,0.0012708743,0.0012480428,0.006025906,0.0006893513,0.001133685,0.0009849314,0.0007994156,0.004689154],"category_scores_gemma":[0.0051726312,0.0002388535,0.0014872788,0.0023458505,0.00031579947,0.000801464,0.0018769464,0.00077951903,0.000918684],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001051448,0.00052686164,0.058011033,0.0007270884,0.0017132348,0.0011449085,0.00029295697,0.09120304,0.035412792,0.007153796,0.02698866,0.7757741],"study_design_scores_gemma":[0.0002100953,0.0004010043,0.028007794,0.000117208074,0.00084690854,0.0011883031,0.00022684463,0.8878384,0.016840996,0.042759467,0.021451874,0.000111123285],"about_ca_topic_score_codex":0.0027561805,"about_ca_topic_score_gemma":0.0043145814,"teacher_disagreement_score":0.006025906,"about_ca_system_score_codex":0.00046703994,"about_ca_system_score_gemma":0.0009953257,"threshold_uncertainty_score":0.01568681},"labels":[],"label_agreement":null},{"id":"W3207474301","doi":"10.1109/ichi52183.2021.00016","title":"A Framework To Build A Causal Knowledge Graph for Chronic Diseases and Cancers By Discovering Semantic Associations from Biomedical Literature","year":2021,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Computer science; Breast cancer; Causality (physics); Word embedding; Semantics (computer science); Disease; Knowledge extraction; Semantic Web; Artificial intelligence; Data science; Machine learning; Information retrieval; Cancer; Embedding; Medicine","score_opus":0.009112322521432663,"score_gpt":0.2980494888575245,"score_spread":0.28893716633609184,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207474301","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0032869405,0.0012441962,0.98355585,0.0015820827,0.000079994854,0.00029978494,0.004434827,0.0016350065,0.0038814538],"genre_scores_gemma":[0.034542,0.0011526706,0.95700085,0.00023564417,0.00005910514,0.000320624,0.005436218,0.00009783713,0.0011551081],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99842525,0.00044797608,0.00021400194,0.0005145884,0.0003383525,0.000059874175],"domain_scores_gemma":[0.99581075,0.0023181858,0.00048121385,0.0005076643,0.0007129654,0.00016922467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0026378427,0.0010611478,0.00069050165,0.017218828,0.0014866734,0.0029600258,0.0017102471,0.0012578161,0.005173083],"category_scores_gemma":[0.009499123,0.0006036634,0.0030927204,0.011149128,0.0011946446,0.004490119,0.002398233,0.001270032,0.0016258135],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010444044,0.00020403604,0.008016128,0.0026800768,0.0005346593,0.0011755525,0.0019308047,0.053973094,0.0048659286,0.46900335,0.018900368,0.4386117],"study_design_scores_gemma":[0.000048042228,0.00009747353,0.002633195,0.0008220271,0.00043429286,0.0010789748,0.000739683,0.17521964,0.0029103605,0.643775,0.1721472,0.00009414458],"about_ca_topic_score_codex":0.0141572775,"about_ca_topic_score_gemma":0.025295492,"teacher_disagreement_score":0.017218828,"about_ca_system_score_codex":0.0015981514,"about_ca_system_score_gemma":0.004586245,"threshold_uncertainty_score":0.028149724},"labels":[],"label_agreement":null},{"id":"W3207996961","doi":"","title":"Expressions, Utterances and Directive Slots","year":2021,"lang":"fr","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Directive; Computer science; Programming language","score_opus":0.012148444604742802,"score_gpt":0.24158996077479833,"score_spread":0.22944151617005554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3207996961","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06265709,0.0020810615,0.6910649,0.0047377255,0.001483568,0.0005401426,0.023052048,0.017304575,0.19707894],"genre_scores_gemma":[0.6212093,0.0015605235,0.25941855,0.0008682698,0.000559428,0.0010417459,0.023656216,0.0061619333,0.08552405],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99719065,0.0011872943,0.00021685776,0.0006078923,0.0005217649,0.0002755551],"domain_scores_gemma":[0.9966509,0.002350208,0.00015444835,0.00034259364,0.00040295895,0.0000988923],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020899323,0.0016266946,0.0006463256,0.0015394145,0.0015767637,0.0035068046,0.0010246553,0.0017313334,0.023683725],"category_scores_gemma":[0.0066474513,0.00106902,0.00063922984,0.0019045551,0.0022248905,0.0060523245,0.0025445665,0.0019361768,0.008834995],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010999275,0.0000894427,0.0022514507,0.00085955806,0.0000503108,0.0016597728,0.019027494,0.0015064697,0.018967077,0.7540917,0.083013386,0.1173834],"study_design_scores_gemma":[0.00014421206,0.0001066838,0.0033144758,0.00032726288,0.00010570822,0.0026126434,0.009460635,0.016471105,0.02272768,0.2918415,0.6527207,0.00016732911],"about_ca_topic_score_codex":0.002305112,"about_ca_topic_score_gemma":0.0012798914,"teacher_disagreement_score":0.023683725,"about_ca_system_score_codex":0.0014080596,"about_ca_system_score_gemma":0.0009613663,"threshold_uncertainty_score":0.07922995},"labels":[],"label_agreement":null},{"id":"W3208248724","doi":"10.5281/zenodo.5645679","title":"Refinement of the COHESIVE Information System towards a unified ontology of food terms for the public health organizations","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Larus Technologies (Canada)","funders":"Ministero della Salute; European Commission","keywords":"Ontology; Public health; Computer science; Knowledge management; Business; Data science; Medicine; Epistemology","score_opus":0.0425059027417626,"score_gpt":0.25731881156820186,"score_spread":0.21481290882643925,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208248724","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019780701,0.00042341914,0.9649921,0.0010326677,0.00011444447,0.000777561,0.004975816,0.0054005845,0.0025027252],"genre_scores_gemma":[0.051177528,0.0002802896,0.9306577,0.00030987264,0.000044175227,0.00040158286,0.015357081,0.00067080033,0.0011009616],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99002236,0.0022752993,0.0022823948,0.002145795,0.0028693422,0.00040488222],"domain_scores_gemma":[0.9772589,0.006169795,0.001389323,0.0069724107,0.0074132713,0.00079625455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011687781,0.0011188698,0.0019109442,0.011656422,0.0029968852,0.0073213987,0.0033017069,0.0019886987,0.003660629],"category_scores_gemma":[0.030883104,0.0013280072,0.0047229864,0.007344092,0.0016592076,0.008424576,0.00803981,0.0035562736,0.001953717],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062243536,0.00089796854,0.0184843,0.0032423937,0.0013302907,0.0013639068,0.011628347,0.026370358,0.025070589,0.3515802,0.035035748,0.5243735],"study_design_scores_gemma":[0.00019315403,0.00035034565,0.01177963,0.0024901875,0.002520923,0.0011782068,0.003531162,0.29862586,0.03709548,0.30424312,0.33768383,0.0003081131],"about_ca_topic_score_codex":0.022420358,"about_ca_topic_score_gemma":0.023846209,"teacher_disagreement_score":0.022420358,"about_ca_system_score_codex":0.0025663844,"about_ca_system_score_gemma":0.010204453,"threshold_uncertainty_score":0.061811686},"labels":[],"label_agreement":null},{"id":"W3208354800","doi":"10.1101/2021.06.01.446587","title":"OBO Foundry in 2021: Operationalizing Open Data Principles to Evaluate Ontologies","year":2021,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"National Institutes of Health","keywords":"Computer science; Ontology; Operationalization; Suite; Interoperability; Data science; Metadata; Authentication (law); World Wide Web; Software engineering; Knowledge management","score_opus":0.08950262238526517,"score_gpt":0.3293741242944935,"score_spread":0.2398715019092283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3208354800","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1091888,0.0003797795,0.8366034,0.010604359,0.0010629048,0.0023914801,0.003091715,0.015084747,0.02159275],"genre_scores_gemma":[0.25837955,0.00015022059,0.7242778,0.0016043575,0.00012230984,0.0012730206,0.0068062306,0.0037341067,0.003652408],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8566031,0.075956106,0.013749,0.009952107,0.038754452,0.0049851853],"domain_scores_gemma":[0.6756917,0.112752765,0.019610815,0.10378895,0.08139204,0.0067637116],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.16850027,0.0015368258,0.0021178417,0.0076214443,0.004184628,0.0205828,0.0047656484,0.0035955356,0.002991085],"category_scores_gemma":[0.28439516,0.0014466158,0.0024541856,0.004005989,0.009301642,0.017304676,0.017996378,0.0073883887,0.0014361133],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013897498,0.001185052,0.049756113,0.0009810139,0.00041008508,0.0008960781,0.010258309,0.051908627,0.014287244,0.62925255,0.03836045,0.20131466],"study_design_scores_gemma":[0.00026719808,0.0005203309,0.010666226,0.0017456432,0.00015370123,0.00034790483,0.0048204935,0.2555118,0.05136377,0.5082746,0.16586916,0.00045918778],"about_ca_topic_score_codex":0.013918974,"about_ca_topic_score_gemma":0.009970491,"teacher_disagreement_score":0.9952344,"about_ca_system_score_codex":0.008690919,"about_ca_system_score_gemma":0.015819216,"threshold_uncertainty_score":0.89112526},"labels":[],"label_agreement":null},{"id":"W3209668695","doi":"10.23977/jaip.2020.040105","title":"Research on Entity Recognition and Knowledge Graph Construction Based on Tcm Medical Records","year":2021,"lang":"en","type":"article","venue":"Journal of Artificial Intelligence Practice","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Knowledge extraction; Ambiguity; Information retrieval; Visualization; Graph; Conditional random field; Artificial intelligence; Data science; Data mining","score_opus":0.1629362907974108,"score_gpt":0.4410426488735801,"score_spread":0.2781063580761693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3209668695","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029765598,0.0014931012,0.9613385,0.00077929214,0.000101839556,0.00016615633,0.0010699022,0.0012662738,0.0040193065],"genre_scores_gemma":[0.4444378,0.0046755574,0.5359733,0.0003174694,0.0001324491,0.000279611,0.0068498137,0.00016625327,0.007167697],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987301,0.00022043745,0.00014374909,0.0005060499,0.0003237848,0.00007582807],"domain_scores_gemma":[0.9986823,0.0005426259,0.0001552313,0.00018910311,0.00038040357,0.00005037138],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009484158,0.000563537,0.00059434766,0.003769376,0.00085000263,0.0014332646,0.0014001916,0.0006658911,0.0022700238],"category_scores_gemma":[0.0041025546,0.00037879765,0.0014636242,0.0050204806,0.00061520125,0.005702992,0.0010516429,0.0007521481,0.0005260191],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000102628524,0.0001541917,0.014889164,0.00095914176,0.00025161062,0.0005505711,0.0008539269,0.092781164,0.006634037,0.087778114,0.00979711,0.78524846],"study_design_scores_gemma":[0.000021686043,0.00006847229,0.007246079,0.00014214724,0.000252377,0.00061775785,0.00040661913,0.8984666,0.010140923,0.05830303,0.024242297,0.00009195484],"about_ca_topic_score_codex":0.015946036,"about_ca_topic_score_gemma":0.012213337,"teacher_disagreement_score":0.015946036,"about_ca_system_score_codex":0.0010688168,"about_ca_system_score_gemma":0.0018327127,"threshold_uncertainty_score":0.031706452},"labels":[],"label_agreement":null},{"id":"W3210285215","doi":"10.3233/sw-233458","title":"OBO Foundry food ontology interconnectivity","year":2024,"lang":"en","type":"article","venue":"Semantic Web","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Genome Canada; U.S. Department of Agriculture; National Science Foundation","keywords":"Ontology; Vocabulary; Interoperability; Computer science; Food industry; SNOMED CT; Government (linguistics); Sustainability; Food processing; Open Biomedical Ontologies; Data science; Knowledge management; World Wide Web; Business; Process ontology; Suggested Upper Merged Ontology; Domain knowledge; Political science; Terminology","score_opus":0.016277562294126212,"score_gpt":0.27359890744408066,"score_spread":0.2573213451499545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210285215","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036155827,0.002737452,0.6974867,0.01250561,0.0019404409,0.0013279191,0.027559945,0.018901365,0.20138477],"genre_scores_gemma":[0.17311513,0.0041537546,0.61520976,0.004529694,0.0005963152,0.0016352887,0.1247072,0.011105808,0.06494706],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9947402,0.0009840801,0.0009883264,0.0010895547,0.0017261544,0.00047166945],"domain_scores_gemma":[0.99016273,0.0019485682,0.00090287405,0.0038790044,0.002404659,0.0007022505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009856704,0.00063488923,0.0009033616,0.006467294,0.0029978144,0.006103914,0.0018246043,0.0015482589,0.007642503],"category_scores_gemma":[0.017084815,0.00080167473,0.0012817438,0.006681192,0.002464276,0.012892103,0.009600493,0.0023644462,0.003668291],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027700237,0.00015339421,0.008204433,0.0009814585,0.000110876885,0.0009977598,0.006867493,0.0014747796,0.012101894,0.65601116,0.09723293,0.21558678],"study_design_scores_gemma":[0.000014628321,0.000018629467,0.0017229096,0.00035591994,0.000036319427,0.00028228568,0.0007007487,0.001732932,0.002601851,0.05233949,0.9401662,0.000028062428],"about_ca_topic_score_codex":0.023242647,"about_ca_topic_score_gemma":0.021020064,"teacher_disagreement_score":0.023242647,"about_ca_system_score_codex":0.0046574413,"about_ca_system_score_gemma":0.009257537,"threshold_uncertainty_score":0.052127838},"labels":[],"label_agreement":null},{"id":"W3210894172","doi":"","title":"THE LINGUISTIC ASPECT OF MEDICAL RESEARCH PAPERS- NEED TO DEVELOP INSIGHT","year":2008,"lang":"en","type":"article","venue":"DOAJ (DOAJ: Directory of Open Access Journals)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Linguistics; Computer science; Value (mathematics); Urdu; Order (exchange); English language; International communication; Medical education; Psychology; Medicine; Business; Communication","score_opus":0.3266894335884901,"score_gpt":0.57767931828221,"score_spread":0.25098988469371986,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3210894172","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03418005,0.036852825,0.1597184,0.4735227,0.011182587,0.001099837,0.0015862602,0.00089000445,0.28096732],"genre_scores_gemma":[0.64982826,0.029308576,0.21455057,0.046436366,0.011757573,0.0014031978,0.0016304965,0.0010084037,0.04407657],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9068234,0.056844156,0.012495909,0.002785956,0.019946558,0.0011040743],"domain_scores_gemma":[0.70391226,0.19939715,0.027807761,0.019980129,0.044455115,0.004447708],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07550277,0.0004948508,0.00091305433,0.010236498,0.003287639,0.021583565,0.00201306,0.002774234,0.00782163],"category_scores_gemma":[0.20065083,0.00059631857,0.00074160413,0.010196,0.010153932,0.026867148,0.0063582845,0.0029871985,0.0040047793],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015029307,0.000082857274,0.005862291,0.005575963,0.00014838096,0.0017308766,0.064363465,0.00034899023,0.00403607,0.68168706,0.058447238,0.17756642],"study_design_scores_gemma":[0.000035151905,0.000051227395,0.0044042864,0.0033698543,0.00010657356,0.003205059,0.031893473,0.00077845773,0.0015975733,0.40113327,0.5533248,0.00010026795],"about_ca_topic_score_codex":0.0010996169,"about_ca_topic_score_gemma":0.0011602085,"teacher_disagreement_score":0.92449725,"about_ca_system_score_codex":0.006115592,"about_ca_system_score_gemma":0.0083045475,"threshold_uncertainty_score":0.3993016},"labels":[],"label_agreement":null},{"id":"W3211905090","doi":"10.1016/j.jcjd.2021.09.065","title":"Data on Patient Record Trajectory for Linkage (DataPRinT Linkage)","year":2021,"lang":"en","type":"article","venue":"Canadian Journal of Diabetes","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Toronto Public Health; North York General Hospital","funders":"","keywords":"Linkage (software); Medicine; Record linkage; Health records; Variety (cybernetics); Medical record; Trajectory; Medical emergency; Genetics; Environmental health; Health care; Internal medicine; Gene; Computer science; Artificial intelligence","score_opus":0.03695076323179959,"score_gpt":0.2672872227547012,"score_spread":0.2303364595229016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3211905090","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013001575,0.00036154856,0.039665718,0.00067226705,0.00027513978,0.0006624681,0.93100226,0.0078688115,0.00649025],"genre_scores_gemma":[0.064931035,0.00060688634,0.08817925,0.00030859126,0.000097301105,0.0016307594,0.837193,0.00085461,0.0061985375],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.995226,0.00058803556,0.0010464723,0.0013524407,0.0013955899,0.00039141462],"domain_scores_gemma":[0.98683053,0.0038650658,0.00179432,0.0041495175,0.0028225589,0.00053800415],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005443696,0.00086507056,0.0010596804,0.009102796,0.0010300835,0.0028284502,0.0010297225,0.0011308403,0.040778756],"category_scores_gemma":[0.029131135,0.00048468038,0.0014569944,0.011624368,0.00025959522,0.0019761822,0.0028269228,0.001148622,0.018362416],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019267204,0.00045921118,0.08621192,0.0045375824,0.00084786414,0.0010425317,0.0019737624,0.006648105,0.011218049,0.019893244,0.47640386,0.3888372],"study_design_scores_gemma":[0.0002785933,0.00029750657,0.08703418,0.000906399,0.0005291251,0.0012077434,0.0011822212,0.010103259,0.028508147,0.016028475,0.85373837,0.00018597078],"about_ca_topic_score_codex":0.018789537,"about_ca_topic_score_gemma":0.016251087,"teacher_disagreement_score":0.040778756,"about_ca_system_score_codex":0.0016579317,"about_ca_system_score_gemma":0.008630638,"threshold_uncertainty_score":0.13641852},"labels":[],"label_agreement":null},{"id":"W3214176644","doi":"10.1016/j.xgen.2021.100028","title":"The Data Use Ontology to streamline responsible access to human biomedical datasets","year":2021,"lang":"en","type":"article","venue":"Cell Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":63,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Google (Canada); Montreal Neurological Institute and Hospital; University Health Network; McGill University; Canada's Michael Smith Genome Sciences Centre; McGill Genome Centre","funders":"National Health and Medical Research Council; Horizon 2020 Framework Programme; National Institutes of Health; ZonMw; FP7 Coherent Development of Research Policies; University of Michigan; Government of the United Kingdom; European Bioinformatics Institute; McGill University; Horizon 2020; EOSC-Life; Bayer; Japan Agency for Medical Research and Development; Novartis; European Commission; Broad Institute; International Business Machines Corporation; National Human Genome Research Institute; Wellcome Trust; Intel Corporation","keywords":"Computer science; Ontology; Data science; Data access; Data sharing; Data management; Data discovery; Metadata; Information retrieval; World Wide Web; Data mining; Database","score_opus":0.08713909626843905,"score_gpt":0.3653044782949046,"score_spread":0.27816538202646557,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214176644","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005642877,0.0005740579,0.8924962,0.009876907,0.00079118856,0.0020759983,0.03530548,0.014360814,0.038876362],"genre_scores_gemma":[0.03410628,0.0013760779,0.8616058,0.005050913,0.00036411043,0.0024620045,0.07849028,0.004242786,0.01230172],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98062253,0.0052210726,0.0048501203,0.0024076928,0.005917811,0.0009807098],"domain_scores_gemma":[0.96042484,0.011315495,0.0026185282,0.016386012,0.007393816,0.0018612485],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0324927,0.0008502031,0.001064491,0.0065282276,0.0028476936,0.007838904,0.0030856612,0.0026957314,0.0053560184],"category_scores_gemma":[0.04189427,0.0011948205,0.0024678914,0.007990186,0.00365215,0.017050814,0.0102630155,0.005186076,0.0056277057],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021780109,0.00024429438,0.0068630762,0.0014020737,0.00015359174,0.00053262257,0.004171466,0.0031097138,0.009442001,0.66503364,0.16446817,0.1443616],"study_design_scores_gemma":[0.000032545995,0.000030381141,0.002306902,0.00065138756,0.000056565554,0.00030234412,0.0006537235,0.0049655107,0.0050243633,0.07676307,0.90913,0.00008325111],"about_ca_topic_score_codex":0.020113427,"about_ca_topic_score_gemma":0.016437596,"teacher_disagreement_score":0.9969143,"about_ca_system_score_codex":0.0047474746,"about_ca_system_score_gemma":0.022811137,"threshold_uncertainty_score":0.17183983},"labels":[],"label_agreement":null},{"id":"W3214230530","doi":"10.3390/ijerph182212025","title":"The Prescription of Drug Ontology 2.0 (PDRO): More Than the Sum of Its Parts","year":2021,"lang":"en","type":"article","venue":"International Journal of Environmental Research and Public Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"Canadian Institutes of Health Research","keywords":"Ontology; Medical prescription; Observational study; Pharmacy; Information system; General partnership; Product (mathematics); Medicine; Health care; Computer science; Data science; Knowledge management; Medical emergency; Business; Family medicine; Nursing; Engineering; Political science","score_opus":0.05631574425175256,"score_gpt":0.37678444851848014,"score_spread":0.3204687042667276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214230530","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017657213,0.005591925,0.8146748,0.031849455,0.0016854047,0.0008324315,0.025232775,0.013191097,0.08928489],"genre_scores_gemma":[0.0969186,0.008618721,0.819246,0.007931148,0.0005361196,0.0006721359,0.045484155,0.0034300236,0.017163044],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963084,0.00084416394,0.000741982,0.00047037005,0.0014299079,0.00020525747],"domain_scores_gemma":[0.99220544,0.0034174295,0.0010428436,0.0017103229,0.001102933,0.0005209499],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005775557,0.000676117,0.0007686301,0.0050749374,0.0016590181,0.006152718,0.001300522,0.002308704,0.0035798943],"category_scores_gemma":[0.010878931,0.0007201147,0.0010577721,0.007491869,0.0024102093,0.008501281,0.0046955454,0.0029370121,0.0020530869],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001408066,0.00015036103,0.0057024527,0.0024341997,0.00014927829,0.0010505625,0.0029356356,0.003748472,0.006933653,0.6122153,0.1017225,0.26281682],"study_design_scores_gemma":[0.000012344945,0.000025713724,0.0019016102,0.00058370124,0.00004904852,0.0010735894,0.0005896891,0.0055384757,0.0014574925,0.06824869,0.9204638,0.000055978646],"about_ca_topic_score_codex":0.009170676,"about_ca_topic_score_gemma":0.01015245,"teacher_disagreement_score":0.009170676,"about_ca_system_score_codex":0.0027914708,"about_ca_system_score_gemma":0.0070200604,"threshold_uncertainty_score":0.0305444},"labels":[],"label_agreement":null},{"id":"W3214982152","doi":"10.11575/prism/39385","title":"Computational Drug Repositioning Based on Integrated Similarity Measures and Deep Learning","year":2020,"lang":"en","type":"dissertation","venue":"PRISM (University of Calgary)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alberta Innovates; Alberta Innovates - Technology Futures; Government of Alberta","keywords":"Deep learning; Drug repositioning; Artificial intelligence; Similarity (geometry); Computer science; Machine learning; Data science; Drug; Medicine; Pharmacology","score_opus":0.008101568733795243,"score_gpt":0.21149121123636994,"score_spread":0.20338964250257469,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3214982152","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.038789302,0.0019089232,0.9531337,0.00062567514,0.00009896532,0.00011434635,0.00022284553,0.0009772185,0.004129039],"genre_scores_gemma":[0.5804136,0.002012679,0.41097346,0.00047508077,0.00020093848,0.00042654513,0.0009626233,0.00018869837,0.0043464536],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995639,0.000114292314,0.000044663368,0.0000879731,0.00012879896,0.000060298047],"domain_scores_gemma":[0.9988223,0.00074444537,0.00010093572,0.00010444471,0.00016393127,0.00006389657],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009796629,0.00077591476,0.0018002036,0.0015251167,0.00044337267,0.0014025811,0.0019787871,0.0012294948,0.0024067583],"category_scores_gemma":[0.0027524594,0.0006549228,0.0012794086,0.0015043082,0.00073418446,0.001641275,0.001688523,0.0018659543,0.0003506804],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000050518534,0.00007281047,0.00065118854,0.0000930739,0.000072997915,0.00004166311,0.000021775068,0.9205434,0.0005370989,0.010296669,0.00094041124,0.06667842],"study_design_scores_gemma":[0.0000034098246,0.000008923213,0.000027121181,0.0000028391423,0.0000038157677,0.00000408263,0.0000016250424,0.9974954,0.00014173142,0.0021939906,0.00011535614,0.0000016946342],"about_ca_topic_score_codex":0.009291862,"about_ca_topic_score_gemma":0.008850931,"teacher_disagreement_score":0.009291862,"about_ca_system_score_codex":0.0014608115,"about_ca_system_score_gemma":0.0021103956,"threshold_uncertainty_score":0.018475533},"labels":[],"label_agreement":null},{"id":"W3216662973","doi":"10.1287/orsc.2021.1524","title":"Learning by Connecting: How Rule Networks Evolve Through Discovery of Relevance","year":2021,"lang":"en","type":"article","venue":"Organization Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Waterloo; University of British Columbia","funders":"","keywords":"Relevance (law); Context (archaeology); Knowledge management; Citation; Process (computing); Computer science; Data science; Political science; World Wide Web","score_opus":0.009245199295435673,"score_gpt":0.25175200045364254,"score_spread":0.24250680115820686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3216662973","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.36963674,0.0014348403,0.5938022,0.005976771,0.00013807746,0.00031330186,0.00067013846,0.00083208823,0.027195903],"genre_scores_gemma":[0.85224557,0.0007384588,0.14248419,0.00036676702,0.00007952305,0.0002638106,0.0006988428,0.00016170123,0.0029611094],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99158317,0.00397255,0.00045999372,0.0023716586,0.0011314643,0.00048127415],"domain_scores_gemma":[0.88273305,0.091387495,0.008154724,0.0109927105,0.0047095073,0.0020225178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010492037,0.00050157955,0.0008850532,0.0050369976,0.0024023603,0.007714062,0.0027629067,0.0025551077,0.0046712463],"category_scores_gemma":[0.12238405,0.0008097644,0.0014800464,0.0035616613,0.0076300506,0.01719433,0.0041103824,0.0023447424,0.0008591396],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021865625,0.0002637802,0.07262943,0.0006446513,0.00037682697,0.0016194681,0.025862992,0.10877417,0.004191809,0.5389833,0.004625965,0.241809],"study_design_scores_gemma":[0.00004675016,0.000074275755,0.0073831887,0.00016386711,0.00013230911,0.00054250605,0.0025903853,0.15066464,0.0016512264,0.82429326,0.01236412,0.00009345436],"about_ca_topic_score_codex":0.0068369494,"about_ca_topic_score_gemma":0.005188421,"teacher_disagreement_score":0.010492037,"about_ca_system_score_codex":0.002196745,"about_ca_system_score_gemma":0.0021760375,"threshold_uncertainty_score":0.05548787},"labels":[],"label_agreement":null},{"id":"W3217105868","doi":"10.3389/fninf.2021.622951","title":"Magnetic Resonance Imaging Sequence Identification Using a Metadata Learning Approach","year":2021,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Heart and Stroke Foundation; Centre for Addiction and Mental Health; Queen's University; St. Joseph’s Healthcare Hamilton; Western University; University of Toronto; University of British Columbia; Hospital for Sick Children; University of Calgary; University Health Network; Health Sciences Centre; St. Michael's Hospital; Sunnybrook Health Science Centre; McMaster University; Robarts Clinical Trials; Indoc Research; Holland Bloorview Kids Rehabilitation Hospital; Baycrest Hospital","funders":"Faculty of Health Sciences, Queen's University; Natural Sciences and Engineering Research Council of Canada; Temerty Family Foundation; H. Lundbeck A/S; Servier; University of British Columbia; London Health Sciences Foundation; Government of Ontario; University of Ottawa; Hospital for Sick Children; Pfizer; Ontario Brain Institute; University of Calgary; Queen's University; Canadian Institutes of Health Research; Centre for Addiction and Mental Health Foundation; McMaster University; Bristol-Myers Squibb","keywords":"Computer science; Metadata; Identification (biology); Artificial intelligence; Magnetic resonance imaging; Sequence (biology); Machine learning; A priori and a posteriori; Software; Information retrieval; Data mining; World Wide Web; Medicine; Radiology","score_opus":0.025523992380939643,"score_gpt":0.26622907274341706,"score_spread":0.24070508036247742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3217105868","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020925267,0.0005549612,0.9682755,0.0007253133,0.00012000219,0.00042506694,0.0025314053,0.0042505297,0.0021920106],"genre_scores_gemma":[0.14076032,0.0005100629,0.84357315,0.0003563327,0.00017944278,0.00043514103,0.011198139,0.0001840959,0.0028033138],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970822,0.00052072946,0.0004883357,0.00093951006,0.00080139295,0.00016794281],"domain_scores_gemma":[0.9944113,0.0018903875,0.00087514066,0.0009873507,0.0016092616,0.00022664799],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032067178,0.0007964836,0.0008089226,0.008701532,0.0013240902,0.001966667,0.001947175,0.0015495281,0.0014556354],"category_scores_gemma":[0.0094808,0.00035367283,0.0017429837,0.004839281,0.0008780675,0.0040547308,0.0020473942,0.0015144928,0.001958078],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038127202,0.00082580914,0.026993167,0.0004695909,0.00015993917,0.00065526227,0.0005369029,0.031368043,0.01875754,0.016643308,0.014771655,0.88843757],"study_design_scores_gemma":[0.00007232859,0.000322595,0.009109878,0.00023622994,0.0002200081,0.0017012027,0.0010470495,0.85340476,0.028984034,0.0674068,0.037340008,0.00015514204],"about_ca_topic_score_codex":0.007133438,"about_ca_topic_score_gemma":0.011355278,"teacher_disagreement_score":0.008701532,"about_ca_system_score_codex":0.0015216288,"about_ca_system_score_gemma":0.003971074,"threshold_uncertainty_score":0.016958952},"labels":[],"label_agreement":null},{"id":"W32353242","doi":"10.1021/acs.langmuir.0c00128","title":"Interview met mevrouw A. de Vries-de Vries","year":2003,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Alberta Innovates; University of Alberta; Government of Canada","keywords":"Geology","score_opus":0.023434818026596062,"score_gpt":0.28877678218029895,"score_spread":0.26534196415370287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W32353242","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022257216,0.07515433,0.006138824,0.6295441,0.13613796,0.00040808012,0.003814741,0.0009767297,0.12556805],"genre_scores_gemma":[0.1388326,0.04767555,0.0040643187,0.10126323,0.027620543,0.00060912274,0.001545361,0.00080561126,0.67758363],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99946386,0.00019665231,0.000028001246,0.0001219571,0.00014323909,0.000046377052],"domain_scores_gemma":[0.9989114,0.00040829275,0.00007608844,0.00002706355,0.0003535751,0.0002235724],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008829647,0.0005888388,0.00063716,0.0005513642,0.00095265906,0.0011103944,0.0005116099,0.0013457823,0.08033288],"category_scores_gemma":[0.0047130426,0.00020683245,0.000158249,0.00035731366,0.00042303282,0.0012180244,0.0008521703,0.0025004654,0.028961286],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008293142,0.000055116885,0.0005507596,0.0002458717,0.0000056098556,0.0005048092,0.0010495405,0.00014222242,0.0008066987,0.0024488312,0.94572943,0.048378237],"study_design_scores_gemma":[0.000004641334,0.00002629763,0.0006941227,0.000117313815,0.0000017601973,0.0006670321,0.0007910252,0.0001398778,0.0001921329,0.0004251044,0.9969313,0.000009312517],"about_ca_topic_score_codex":0.0017504753,"about_ca_topic_score_gemma":0.0015935428,"teacher_disagreement_score":0.08033288,"about_ca_system_score_codex":0.0007668801,"about_ca_system_score_gemma":0.0006697678,"threshold_uncertainty_score":0.26874024},"labels":[],"label_agreement":null},{"id":"W400800649","doi":"","title":"An actor network approach to the study of scientific evidence in Canada","year":2006,"lang":"en","type":"article","venue":"Medical Entomology and Zoology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.022465669702541977,"score_gpt":0.27566565902907314,"score_spread":0.25319998932653115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W400800649","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40203065,0.013284976,0.37991065,0.047843933,0.00025680728,0.0013007524,0.025959792,0.0008051124,0.12860738],"genre_scores_gemma":[0.85007817,0.0036299927,0.13232534,0.00036158183,0.000055707387,0.0002505137,0.002902741,0.00006299075,0.010332817],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9964244,0.00170675,0.00029539946,0.0005540256,0.0007187084,0.00030065826],"domain_scores_gemma":[0.96472865,0.027468467,0.0016851121,0.0006618248,0.004313134,0.0011427163],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.006033059,0.00035494255,0.0008270717,0.011385118,0.004603769,0.005940287,0.0015314537,0.0010873998,0.005671626],"category_scores_gemma":[0.033153594,0.00047246338,0.0007333286,0.018998947,0.002855525,0.003269175,0.0021443844,0.0011899235,0.00023251674],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022468875,0.000085757965,0.054251865,0.0006681223,0.0003898814,0.00088344683,0.0057603023,0.07411216,0.0005082196,0.76019907,0.010865939,0.09205062],"study_design_scores_gemma":[0.00012861301,0.000041713138,0.04311458,0.0007368742,0.00061847974,0.00032699484,0.011434034,0.38075185,0.0010598998,0.4220887,0.13956907,0.00012925769],"about_ca_topic_score_codex":0.9598626,"about_ca_topic_score_gemma":0.9609836,"teacher_disagreement_score":0.99539626,"about_ca_system_score_codex":0.043326117,"about_ca_system_score_gemma":0.06275207,"threshold_uncertainty_score":0.31435448},"labels":[],"label_agreement":null},{"id":"W4200088257","doi":"10.7554/elife.73430.sa2","title":"Author response: Experiments from unfinished Registered Reports in the Reproducibility Project: Cancer Biology","year":2021,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kinexus Bioinformatics Corporation (Canada); University of British Columbia; Applied Biological Materials (Canada)","funders":"","keywords":"Reproducibility; Biology; Medical physics; Medicine; Statistics; Mathematics","score_opus":0.15642684775999868,"score_gpt":0.4516190631253949,"score_spread":0.2951922153653962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200088257","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008001369,0.0013445547,0.012396356,0.5792747,0.25616202,0.003665149,0.055103518,0.006411796,0.07764051],"genre_scores_gemma":[0.063638486,0.0021435258,0.031497784,0.32861757,0.026349116,0.009909543,0.053985402,0.01034234,0.4735163],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.92263746,0.024675224,0.0059409216,0.004457058,0.039428715,0.0028605913],"domain_scores_gemma":[0.50604975,0.1666353,0.009834377,0.055188563,0.25040898,0.0118830195],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.061894834,0.0009255163,0.0018700064,0.002339134,0.004818358,0.0056854025,0.0026975358,0.0068500033,0.20076114],"category_scores_gemma":[0.3689779,0.00078923884,0.0015725213,0.0035382712,0.0026183745,0.0030844852,0.004946183,0.007179601,0.095644206],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035959575,0.000046441266,0.00037339516,0.00030278167,0.000032786418,0.00008066885,0.00022226683,0.000035113128,0.00051164406,0.0008306926,0.9900738,0.007130709],"study_design_scores_gemma":[0.00037509165,0.00011702916,0.002869559,0.00029863048,0.00008421139,0.00015890105,0.0007320171,0.0001850513,0.0026487375,0.0023190135,0.99014205,0.00006974715],"about_ca_topic_score_codex":0.0076292343,"about_ca_topic_score_gemma":0.012278606,"teacher_disagreement_score":0.93810517,"about_ca_system_score_codex":0.0034075996,"about_ca_system_score_gemma":0.0164275,"threshold_uncertainty_score":0.6716129},"labels":[],"label_agreement":null},{"id":"W4200167655","doi":"10.1109/ispa-bdcloud-socialcom-sustaincom52081.2021.00208","title":"MediNER: Understanding Diabetes Management Strategies Based on Social Media Discourse","year":2021,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Canada Research Chairs","keywords":"Psychological intervention; Social media; Diabetes mellitus; Diabetes management; Type 2 diabetes; Digital advertising; Medicine; Knowledge management; Computer science; Nursing; World Wide Web","score_opus":0.05013714387694321,"score_gpt":0.3035656591723608,"score_spread":0.25342851529541754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200167655","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.73091775,0.0077147954,0.17070152,0.019439043,0.0004575824,0.00076564157,0.022868434,0.0022983833,0.04483671],"genre_scores_gemma":[0.9053014,0.0023103782,0.08148092,0.0006698894,0.0002143877,0.00029627263,0.0065717637,0.00008474562,0.0030702783],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99929,0.0003250836,0.000074053765,0.00015738377,0.00010821099,0.000045299643],"domain_scores_gemma":[0.9935982,0.0050134985,0.0007394639,0.00025920605,0.00022725819,0.00016230743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019315978,0.00054324337,0.00025490992,0.0048103645,0.000770328,0.0030647526,0.0004718126,0.0008286415,0.0033530802],"category_scores_gemma":[0.007114015,0.00017512296,0.00044453162,0.0021709485,0.0009496857,0.007151343,0.0015244081,0.00070313696,0.00054105127],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008520881,0.00063902437,0.1811674,0.0027551053,0.000349802,0.0015957921,0.10370311,0.004888394,0.025589524,0.12679414,0.021061204,0.5306044],"study_design_scores_gemma":[0.0001004562,0.00035105288,0.23124164,0.0020959955,0.00047745463,0.0021115583,0.121551245,0.14740747,0.018215593,0.1703745,0.30574545,0.00032757223],"about_ca_topic_score_codex":0.0034145326,"about_ca_topic_score_gemma":0.004701803,"teacher_disagreement_score":0.0048103645,"about_ca_system_score_codex":0.000833033,"about_ca_system_score_gemma":0.00065023464,"threshold_uncertainty_score":0.011217177},"labels":[],"label_agreement":null},{"id":"W4200448776","doi":"10.1101/2021.11.12.467727","title":"The Xenopus Phenotype Ontology: bridging model organism phenotype data to human health and development","year":2021,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; National Institute of Child Health and Human Development; National Human Genome Research Institute","keywords":"Ontology; Open Biomedical Ontologies; Computer science; Interoperability; Ontology-based data integration; Phenotype; Process ontology; Computational biology; Bridging (networking); Upper ontology; Suggested Upper Merged Ontology; Biology; Bioinformatics; Information retrieval; World Wide Web; Semantic Web; Gene; Genetics","score_opus":0.04466711721631948,"score_gpt":0.2819049161111436,"score_spread":0.2372377988948241,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4200448776","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023612045,0.000835581,0.91328555,0.0039024893,0.00031949911,0.0004917616,0.019148191,0.010194485,0.028210374],"genre_scores_gemma":[0.16668636,0.0023887127,0.7919816,0.0012134037,0.00013347421,0.0007151835,0.027088728,0.0026592724,0.0071332427],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99857795,0.0004144219,0.000304079,0.000262287,0.00035928527,0.000082033985],"domain_scores_gemma":[0.99729127,0.0009868629,0.00034988965,0.0008007042,0.00040095104,0.00017026634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003100404,0.00048519537,0.00041942374,0.0019278526,0.0006690864,0.0021875624,0.0012316293,0.0007144456,0.004134832],"category_scores_gemma":[0.0038612795,0.0003772029,0.0010178618,0.0017554181,0.0013570717,0.0033086026,0.0020368097,0.0011361089,0.0010418732],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034158857,0.00014502998,0.014291248,0.0026863024,0.00018758682,0.0015506526,0.0034377344,0.011556024,0.034476776,0.6518376,0.07286983,0.20661962],"study_design_scores_gemma":[0.000050070445,0.00007428052,0.0100134695,0.0011288662,0.00014226601,0.001406718,0.0011947978,0.020904783,0.024664573,0.11879838,0.82150674,0.000115162744],"about_ca_topic_score_codex":0.009127561,"about_ca_topic_score_gemma":0.009491202,"teacher_disagreement_score":0.009127561,"about_ca_system_score_codex":0.0014663467,"about_ca_system_score_gemma":0.0030711098,"threshold_uncertainty_score":0.01814884},"labels":[],"label_agreement":null},{"id":"W4205456307","doi":"10.3115/1654415.1654419","title":"Term generalization and synonym resolution for biological abstracts","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"Natural Sciences and Engineering Research Council of Canada; University of Alberta","keywords":"Computer science; Synonym (taxonomy); Information retrieval; Ontology; Task (project management); Generalization; Hierarchy; Gene ontology; Categorization; Text categorization; WordNet; Artificial intelligence; Field (mathematics); Function (biology); Natural language processing; Biology; Gene; Mathematics; Genus","score_opus":0.020299980912479348,"score_gpt":0.2687625671374854,"score_spread":0.24846258622500608,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205456307","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052209392,0.0026344436,0.9299484,0.0012786258,0.00051038567,0.00077638857,0.0034603225,0.005756102,0.0034259288],"genre_scores_gemma":[0.19689113,0.0013814708,0.78499335,0.0003992149,0.00075602543,0.0008572783,0.010661404,0.0005793338,0.0034807364],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9894488,0.0026282216,0.001867891,0.0026582677,0.0029656382,0.00043121082],"domain_scores_gemma":[0.9735859,0.01575505,0.0023058073,0.0042128433,0.0036497456,0.00049067725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011503149,0.0015998618,0.0023158533,0.017847953,0.002567377,0.0027071766,0.002909261,0.0021599785,0.004550777],"category_scores_gemma":[0.043097693,0.0006344271,0.003103132,0.011944114,0.001362385,0.008364159,0.0038154274,0.0031333028,0.0034310694],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005296183,0.00022869333,0.005963189,0.0010516493,0.00027664722,0.00083563646,0.0014666918,0.011494293,0.019874414,0.02026569,0.025632007,0.91238153],"study_design_scores_gemma":[0.00027246794,0.00044159067,0.01447331,0.00047014342,0.00064754125,0.0050708232,0.0020930795,0.6366968,0.03794525,0.23509113,0.066442594,0.00035523163],"about_ca_topic_score_codex":0.0030922072,"about_ca_topic_score_gemma":0.0029078186,"teacher_disagreement_score":0.017847953,"about_ca_system_score_codex":0.0014517425,"about_ca_system_score_gemma":0.0023894792,"threshold_uncertainty_score":0.060835183},"labels":[],"label_agreement":null},{"id":"W4205464263","doi":"10.17504/protocols.io.br8ym9xw","title":"SARS-CoV-2 NCBI submission protocol: SRA, BioSample, and BioProject v2","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Protocol (science); Biology; Coronavirus disease 2019 (COVID-19); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); Metadata; Computational biology; Computer science; World Wide Web; Pathology; Medicine","score_opus":0.053824740934302874,"score_gpt":0.3595427230010354,"score_spread":0.3057179820667325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4205464263","genre_codex":"dataset","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075739482,0.0035959496,0.16980815,0.008838196,0.009463631,0.039548475,0.5888666,0.06727707,0.10502802],"genre_scores_gemma":[0.010580673,0.002386851,0.10562874,0.006326317,0.0014275126,0.036912624,0.7289176,0.024784522,0.083035156],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9789298,0.007386196,0.0039565014,0.0025105712,0.005457831,0.0017590437],"domain_scores_gemma":[0.9696073,0.005962978,0.0020784533,0.008692742,0.011707479,0.0019509937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019729476,0.0023806456,0.0029108312,0.0054219426,0.005403578,0.0057299235,0.005098812,0.004625683,0.2613089],"category_scores_gemma":[0.03969093,0.003465743,0.0019933698,0.0044635804,0.0017513036,0.005188588,0.006559578,0.005367981,0.35880437],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009397487,0.00016835617,0.0009839353,0.0019136395,0.000047480473,0.00018503964,0.0003713515,0.00013181477,0.02258917,0.0032339648,0.94948703,0.019948345],"study_design_scores_gemma":[0.00030086449,0.00030574642,0.002207273,0.0008651298,0.000042691565,0.00036344765,0.0003180115,0.00031335745,0.021122625,0.002579358,0.97144675,0.0001348189],"about_ca_topic_score_codex":0.0039133755,"about_ca_topic_score_gemma":0.0057794736,"teacher_disagreement_score":0.2613089,"about_ca_system_score_codex":0.0021035722,"about_ca_system_score_gemma":0.009147727,"threshold_uncertainty_score":0.8741654},"labels":[],"label_agreement":null},{"id":"W4206103793","doi":"10.2196/29803","title":"Identification of Prediabetes Discussions in Unstructured Clinical Documentation: Validation of a Natural Language Processing Algorithm","year":2021,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Diabetes and Digestive and Kidney Diseases; National Heart, Lung, and Blood Institute; Johns Hopkins University","keywords":"Prediabetes; Documentation; Computer science; Artificial intelligence; Machine learning; Psychological intervention; Natural language processing; Medicine; Nursing","score_opus":0.010383892816191157,"score_gpt":0.35510812110155093,"score_spread":0.34472422828535976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206103793","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.56481683,0.0012097327,0.4062353,0.0024797814,0.00022181141,0.005251172,0.008295676,0.005942988,0.005546764],"genre_scores_gemma":[0.375267,0.00025065022,0.6157276,0.00029356917,0.00005301614,0.0015293161,0.0060193585,0.000095639705,0.00076391274],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.987225,0.0065550557,0.0023844475,0.002011498,0.0015737438,0.0002502427],"domain_scores_gemma":[0.90772825,0.0730242,0.0055952687,0.0029005124,0.010018671,0.0007330386],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01666679,0.00078391255,0.0005702453,0.005009908,0.0011022183,0.0029484357,0.0013923581,0.0014643869,0.0015373578],"category_scores_gemma":[0.060682192,0.00026626684,0.0007212445,0.00212234,0.0008290061,0.0021767945,0.0020616974,0.0011163385,0.0011337061],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002006865,0.0019944196,0.15131731,0.003878522,0.00033309078,0.0019717028,0.012911711,0.019294398,0.04081401,0.00505571,0.015737092,0.7446852],"study_design_scores_gemma":[0.00074433745,0.0011776949,0.076613285,0.001317473,0.00040921156,0.0034122905,0.011123141,0.7943891,0.062996544,0.015826978,0.03174113,0.00024883094],"about_ca_topic_score_codex":0.004089124,"about_ca_topic_score_gemma":0.0040481025,"teacher_disagreement_score":0.01666679,"about_ca_system_score_codex":0.0015487021,"about_ca_system_score_gemma":0.003930771,"threshold_uncertainty_score":0.08814341},"labels":[],"label_agreement":null},{"id":"W4206440575","doi":"10.1007/s12021-021-09557-0","title":"Is Neuroscience FAIR? A Call for Collaborative Standardisation of Neuroscience Data","year":2022,"lang":"en","type":"article","venue":"Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institute of Neurological Disorders and Stroke; National Institute of Mental Health","keywords":"Neuroinformatics; Neuroscience; Open science; Computational neuroscience; Clinical neuroscience; Cognitive science; Computer science; Psychology; Neurology","score_opus":0.07474942070029075,"score_gpt":0.33520591215909257,"score_spread":0.2604564914588018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206440575","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005011005,0.006972461,0.2511925,0.7125934,0.007374063,0.0003532691,0.0001934301,0.0009315038,0.015378292],"genre_scores_gemma":[0.14860521,0.0113964835,0.64903796,0.16743544,0.010325132,0.002111679,0.001897559,0.002037287,0.0071532996],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.7280658,0.14835373,0.032555148,0.02745142,0.053916827,0.009657126],"domain_scores_gemma":[0.4067839,0.28087565,0.022159703,0.17262742,0.09039443,0.027158935],"candidate_categories":["metaresearch","open_science"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.45050368,0.0013654864,0.0045012655,0.009889231,0.011664478,0.052334227,0.014963847,0.017723419,0.0048267227],"category_scores_gemma":[0.40395576,0.0020658919,0.0041671325,0.008242749,0.05839325,0.1208559,0.04652546,0.04371943,0.002448654],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011141218,0.00022578285,0.002511947,0.0008738261,0.00023878462,0.00034581477,0.01844728,0.0018424449,0.001182775,0.78297573,0.059802048,0.13144216],"study_design_scores_gemma":[0.000030766787,0.000040900795,0.0005785834,0.0013002815,0.000042157775,0.00018773608,0.008396702,0.0011338316,0.0005937043,0.8023354,0.18522742,0.00013252375],"about_ca_topic_score_codex":0.009517495,"about_ca_topic_score_gemma":0.0061927433,"teacher_disagreement_score":0.98503613,"about_ca_system_score_codex":0.014998748,"about_ca_system_score_gemma":0.0707779,"threshold_uncertainty_score":0.6776268},"labels":[],"label_agreement":null},{"id":"W4206514926","doi":"10.1109/bibm52615.2021.9669364","title":"Problem Oriented Diagnostic Service for Describing Clinical Cases based on the GraphQL POMR Approach","year":2021,"lang":"en","type":"article","venue":"2021 IEEE International Conference on Bioinformatics and Biomedicine (BIBM)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Computer science; Schema (genetic algorithms); Structuring; Workflow; SOAP; Service (business); Medical diagnosis; Artificial intelligence; Process (computing); Software engineering; World Wide Web; Machine learning; Programming language; Database; Medicine","score_opus":0.14902245842389547,"score_gpt":0.3475170971003448,"score_spread":0.19849463867644931,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206514926","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0054100035,0.000508271,0.8948745,0.0035406286,0.00029774263,0.0010391173,0.0113022495,0.06270497,0.020322476],"genre_scores_gemma":[0.12921861,0.0012848243,0.81734526,0.0025997949,0.00022950178,0.0009308017,0.032653414,0.0030363973,0.012701384],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99683493,0.0008091335,0.0006929058,0.0005297574,0.0009130662,0.00022021457],"domain_scores_gemma":[0.99696404,0.0011501836,0.00023809075,0.0007090832,0.0007156697,0.00022289941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033892116,0.0010245504,0.00093257957,0.0038342662,0.0009714665,0.005360026,0.002411589,0.0017280221,0.011176244],"category_scores_gemma":[0.009271656,0.00045672778,0.002167958,0.0031378262,0.0011150276,0.004461716,0.003912922,0.0016960633,0.0060264342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007163529,0.00041133625,0.005822189,0.0019446814,0.00021926721,0.004086234,0.003335001,0.021575805,0.013361839,0.47465476,0.15045702,0.3234155],"study_design_scores_gemma":[0.0002216336,0.000119335324,0.0016223777,0.00044684645,0.0001489555,0.002937584,0.0013169865,0.20278628,0.012911639,0.23243113,0.54484445,0.00021269865],"about_ca_topic_score_codex":0.01361502,"about_ca_topic_score_gemma":0.008108728,"teacher_disagreement_score":0.01361502,"about_ca_system_score_codex":0.002295746,"about_ca_system_score_gemma":0.0033520355,"threshold_uncertainty_score":0.037388325},"labels":[],"label_agreement":null},{"id":"W4206758744","doi":"10.2139/ssrn.3990520","title":"What’S in a Pseudonym? Adriana Z. Robertson1 and Albert H. Yoon2","year":2021,"lang":"en","type":"article","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Pseudonym; Political science; Philosophy; Theology","score_opus":0.006491625450582175,"score_gpt":0.24381322932075752,"score_spread":0.23732160387017534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206758744","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018383978,0.04809421,0.21593338,0.60883266,0.036218844,0.00015062255,0.0012556736,0.0010161154,0.070114575],"genre_scores_gemma":[0.43763766,0.07303088,0.29953,0.12004756,0.022684677,0.00041026922,0.0041677016,0.0028702123,0.03962109],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9865145,0.007753189,0.0012159436,0.0011570896,0.0030383433,0.0003208893],"domain_scores_gemma":[0.96161467,0.025400944,0.0021783493,0.002786701,0.006508236,0.001511119],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014129904,0.0004738834,0.0010186203,0.0033931034,0.004198969,0.008853784,0.001094849,0.0030420795,0.011026776],"category_scores_gemma":[0.07497003,0.00049426575,0.0008162351,0.0045015328,0.009248869,0.032963026,0.0038520922,0.0052234773,0.0059350207],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002718557,0.000080199185,0.0049037132,0.00086999004,0.000085602805,0.0013032302,0.007993538,0.00030814143,0.002263827,0.5109463,0.23673436,0.2342392],"study_design_scores_gemma":[0.00003162661,0.000031044223,0.0010817744,0.0010222483,0.000042438973,0.0032371383,0.00545222,0.0016691717,0.0011727398,0.42657793,0.5596019,0.00007985002],"about_ca_topic_score_codex":0.002072489,"about_ca_topic_score_gemma":0.0021204462,"teacher_disagreement_score":0.014129904,"about_ca_system_score_codex":0.0018247006,"about_ca_system_score_gemma":0.0037506693,"threshold_uncertainty_score":0.074727},"labels":[],"label_agreement":null},{"id":"W4206790482","doi":"10.3115/1654415.1654434","title":"BioKI:Enzymes","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Set (abstract data type); Information retrieval; Process (computing); World Wide Web; Programming language","score_opus":0.005581675927201015,"score_gpt":0.22819468154465544,"score_spread":0.22261300561745442,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4206790482","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030398055,0.0077949744,0.09338462,0.0035024136,0.001048167,0.0009922993,0.37038916,0.46521118,0.054637283],"genre_scores_gemma":[0.01523908,0.011511225,0.16687325,0.0025136487,0.00038405217,0.0017542373,0.7125168,0.045941595,0.043266136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974226,0.00042807017,0.00072418735,0.00045491653,0.0007663312,0.0002038947],"domain_scores_gemma":[0.9935894,0.0017392691,0.0011449832,0.0012519658,0.0014816491,0.0007927875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029107013,0.0036352233,0.0028512892,0.0075439885,0.0012340836,0.0073713395,0.0033784069,0.0021826457,0.066091426],"category_scores_gemma":[0.013018236,0.0018787035,0.0021503945,0.010044082,0.00063844107,0.0071981684,0.004039368,0.0040529943,0.18400359],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011572128,0.00014897308,0.0021015324,0.007934729,0.0002906264,0.00056546985,0.00033990404,0.0006112412,0.010690422,0.0109088635,0.8760342,0.08921675],"study_design_scores_gemma":[0.00022314276,0.000053310127,0.0012997699,0.0004385755,0.000107028005,0.0005147565,0.00007798411,0.000930275,0.00834673,0.005872551,0.9820208,0.00011510145],"about_ca_topic_score_codex":0.0017338993,"about_ca_topic_score_gemma":0.0016589061,"teacher_disagreement_score":0.066091426,"about_ca_system_score_codex":0.0015824473,"about_ca_system_score_gemma":0.004942976,"threshold_uncertainty_score":0.22109783},"labels":[],"label_agreement":null},{"id":"W4210303716","doi":"10.18061/bssb.v1i1.8340","title":"Enhanced monography in a collaboratively evolved hub for systematic biology","year":2022,"lang":"en","type":"article","venue":"Bulletin of the Society of Systematic Biologists","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Heritage College","funders":"Ohio State University; National Science Foundation","keywords":"Biology","score_opus":0.022669994667049474,"score_gpt":0.277791874500686,"score_spread":0.2551218798336365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210303716","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021397114,0.0036041052,0.5201482,0.021030288,0.009574779,0.0010633038,0.042756945,0.10614777,0.27427748],"genre_scores_gemma":[0.06623398,0.002148579,0.4913848,0.001906686,0.0020858431,0.00061786314,0.07204615,0.013420077,0.35015604],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988551,0.0003609547,0.00011669046,0.00023726253,0.00035594404,0.00007404329],"domain_scores_gemma":[0.9918168,0.0018595321,0.00026574504,0.0026432055,0.0015750438,0.0018396969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036748322,0.0005207456,0.00046348904,0.00533239,0.0014118046,0.0038379761,0.001417437,0.0013380802,0.14773525],"category_scores_gemma":[0.009101033,0.00043676468,0.0005384532,0.0048625623,0.0005299965,0.006613485,0.0064893914,0.0010754013,0.059073206],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036591536,0.00017894348,0.0021079709,0.00074288866,0.000047210087,0.0005651006,0.001506618,0.0013019886,0.009908019,0.028816683,0.4395676,0.5148911],"study_design_scores_gemma":[0.000028656814,0.00004183802,0.0010855646,0.000120510434,0.000024058316,0.00017921632,0.00019461729,0.0023166928,0.0019153808,0.008231315,0.9858306,0.00003162168],"about_ca_topic_score_codex":0.002425791,"about_ca_topic_score_gemma":0.0061047357,"teacher_disagreement_score":0.14773525,"about_ca_system_score_codex":0.00082450296,"about_ca_system_score_gemma":0.0033851422,"threshold_uncertainty_score":0.49422365},"labels":[],"label_agreement":null},{"id":"W4210372271","doi":"10.2196/preprints.12847","title":"Visibility of Community Nursing Within an Administrative Health Classification System: Evaluation of Content Coverage (Preprint)","year":2018,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Addiction and Mental Health; University of British Columbia","funders":"","keywords":"Psychological intervention; Nursing; Nursing Interventions Classification; Intervention (counseling); Medicine","score_opus":0.33428558292246774,"score_gpt":0.45015940973105306,"score_spread":0.11587382680858532,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210372271","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96642715,0.0008709463,0.007368361,0.0010426656,0.0001788584,0.0085027125,0.004358858,0.00013312993,0.011117355],"genre_scores_gemma":[0.97040695,0.0003665773,0.011505851,0.00021054968,0.00007902137,0.01392491,0.002652364,0.00007072892,0.00078304944],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9018565,0.053804766,0.015213407,0.0038431012,0.023512127,0.0017701673],"domain_scores_gemma":[0.54097706,0.3333375,0.043223836,0.01463474,0.06495162,0.0028752426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10865469,0.00034168392,0.00095872744,0.009048583,0.0017226295,0.0052009434,0.0016408493,0.0007464418,0.005647545],"category_scores_gemma":[0.37439272,0.0004416144,0.0022072606,0.010903577,0.002394586,0.005310563,0.006661184,0.0010140265,0.0006655792],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014157846,0.0008540635,0.7759004,0.004309996,0.00073492684,0.00012964018,0.06009942,0.00065758463,0.00036425501,0.0031487814,0.006026313,0.14635882],"study_design_scores_gemma":[0.00030325723,0.0021510539,0.91429996,0.003183185,0.00056330714,0.00016004004,0.05869094,0.004816517,0.0010919624,0.001761235,0.012844885,0.00013381822],"about_ca_topic_score_codex":0.0072168875,"about_ca_topic_score_gemma":0.006189431,"teacher_disagreement_score":0.10865469,"about_ca_system_score_codex":0.0061270106,"about_ca_system_score_gemma":0.008656596,"threshold_uncertainty_score":0.5746278},"labels":[],"label_agreement":null},{"id":"W4210447837","doi":"10.1093/database/baac001","title":"Authors’ attitude toward adopting a new workflow to improve the computability of phenotype publications","year":2022,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada; University of Ottawa; University of Manitoba","funders":"","keywords":"Workflow; Ambiguity; Computer science; World Wide Web; Knowledge management; Database","score_opus":0.04305262788612411,"score_gpt":0.30935208021017413,"score_spread":0.26629945232405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210447837","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.96861035,0.00047479555,0.013925717,0.007869333,0.00014195198,0.0001683161,0.000093541035,0.0002289712,0.008487062],"genre_scores_gemma":[0.97990304,0.00054706057,0.01591503,0.0014026016,0.00010127009,0.00011365597,0.00013703371,0.00007524936,0.0018049706],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.94992065,0.023392592,0.0073907776,0.002812623,0.0148430625,0.0016402426],"domain_scores_gemma":[0.5851863,0.23606034,0.07675066,0.032346208,0.05456159,0.015094836],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08703584,0.00032245094,0.00034878493,0.0030441063,0.0022882072,0.0081009185,0.0012914626,0.0017985984,0.0020788903],"category_scores_gemma":[0.24119855,0.0005927975,0.0010708985,0.0019332719,0.0032609443,0.005340707,0.0034839758,0.0020158112,0.0008866987],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005558903,0.0004540347,0.66463864,0.0009989887,0.000290583,0.0012171022,0.14873518,0.002587598,0.020390933,0.007319152,0.004580356,0.14823155],"study_design_scores_gemma":[0.00028595235,0.0023758123,0.6059825,0.0017144007,0.0005038195,0.0049431017,0.19437326,0.011897688,0.016660266,0.01703836,0.14324059,0.0009842027],"about_ca_topic_score_codex":0.0022863334,"about_ca_topic_score_gemma":0.0026768567,"teacher_disagreement_score":0.91296417,"about_ca_system_score_codex":0.0024617913,"about_ca_system_score_gemma":0.0050076037,"threshold_uncertainty_score":0.46029502},"labels":[],"label_agreement":null},{"id":"W4210586797","doi":"10.1007/978-981-16-8656-6_5","title":"QL4POMR Interface as a Graph-Based Clinical Diagnosis Web Service","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in operations research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Schema (genetic algorithms); SOAP; Interface (matter); Computer science; World Wide Web; Medical record; Medicine; Information retrieval; Surgery","score_opus":0.10082866818405037,"score_gpt":0.4373734154501951,"score_spread":0.3365447472661447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4210586797","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068952222,0.00037574468,0.5011516,0.0017211074,0.0002915703,0.0006783869,0.03430341,0.4196009,0.034982074],"genre_scores_gemma":[0.18513364,0.0013692969,0.504775,0.0076806126,0.00044901282,0.0017948835,0.13073029,0.07080378,0.097263575],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99935895,0.00017835539,0.000078949925,0.000112878166,0.0002156979,0.00005505158],"domain_scores_gemma":[0.997863,0.001239378,0.00008594063,0.0003267071,0.00031692875,0.00016816604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015797482,0.00095751695,0.0007890137,0.002718952,0.00037150248,0.0025842905,0.0018741122,0.0016302359,0.10268907],"category_scores_gemma":[0.005430481,0.00058749196,0.0010625118,0.0015510212,0.0003668969,0.0015776062,0.0030710688,0.00088401826,0.0354427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002235953,0.00034343087,0.0041142334,0.0015796223,0.00021791221,0.00161025,0.0005206245,0.008439003,0.019694718,0.025335059,0.5320474,0.4038618],"study_design_scores_gemma":[0.00091809395,0.00019729963,0.0037577664,0.00038601848,0.00016751933,0.0015093862,0.00023584708,0.15164474,0.030630967,0.044172253,0.76617515,0.00020488953],"about_ca_topic_score_codex":0.005432222,"about_ca_topic_score_gemma":0.004219379,"teacher_disagreement_score":0.10268907,"about_ca_system_score_codex":0.0011530569,"about_ca_system_score_gemma":0.0009210882,"threshold_uncertainty_score":0.34352916},"labels":[],"label_agreement":null},{"id":"W4213157973","doi":"10.2196/30345","title":"Evaluation of Natural Language Processing for the Identification of Crohn Disease–Related Variables in Spanish Electronic Health Records: A Validation Study for the PREMONITION-CD Project","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Recall; Health records; F1 score; Precision and recall; Natural language processing; Electronic health record; Computer science; Gold standard (test); Medicine; Disease; Artificial intelligence; Health care; Internal medicine; Psychology","score_opus":0.02837736973404331,"score_gpt":0.37198813947420556,"score_spread":0.34361076974016225,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213157973","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97710335,0.00088888913,0.014211543,0.00026948855,0.00008332998,0.0017505636,0.0036782308,0.0006761715,0.0013384463],"genre_scores_gemma":[0.9178334,0.00057522015,0.052704543,0.0003432164,0.0001072263,0.0018387764,0.02536739,0.00014986338,0.0010804388],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9793332,0.013202532,0.0023035123,0.0025426736,0.0022445663,0.00037346198],"domain_scores_gemma":[0.91951585,0.056694876,0.0033441763,0.0050176377,0.014435004,0.0009925145],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02745057,0.001056169,0.00057670625,0.003107683,0.0006103454,0.001609053,0.0011295866,0.0013858435,0.0006969433],"category_scores_gemma":[0.07667047,0.0002682789,0.0012184805,0.0014039704,0.0009517531,0.0012839057,0.0019526725,0.00079227553,0.0005204592],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0066418666,0.012480322,0.48769802,0.0066099027,0.002394658,0.002060039,0.012572178,0.021407425,0.03262273,0.0015127317,0.014801282,0.3991989],"study_design_scores_gemma":[0.00325298,0.0096387025,0.7214372,0.0013479652,0.001928865,0.0026392809,0.0076347436,0.17668635,0.04003398,0.0017814349,0.033312887,0.0003056122],"about_ca_topic_score_codex":0.010973531,"about_ca_topic_score_gemma":0.006558621,"teacher_disagreement_score":0.02745057,"about_ca_system_score_codex":0.0013993464,"about_ca_system_score_gemma":0.0025569396,"threshold_uncertainty_score":0.14517426},"labels":[],"label_agreement":null},{"id":"W4213273087","doi":"10.7326/acpjc-2005-142-1-a08","title":"Finding the gold in MEDLINE: Clinical Queries","year":2005,"lang":"en","type":"article","venue":"ACP Journal Club","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McMaster University; Hamilton Health Sciences","funders":"","keywords":"Medicine; MEDLINE; Gold standard (test); Family medicine; Internal medicine","score_opus":0.044244470189906474,"score_gpt":0.3761877656943085,"score_spread":0.33194329550440205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213273087","genre_codex":"editorial","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007700579,0.05331964,0.0010584177,0.21852161,0.7009035,0.00089392497,0.0014860281,0.00093180564,0.022114996],"genre_scores_gemma":[0.0052405708,0.06609982,0.002054536,0.12066112,0.7559482,0.0006634461,0.0018515978,0.0004882935,0.04699244],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9901578,0.0023718253,0.0029543443,0.0005698643,0.0034913763,0.00045480317],"domain_scores_gemma":[0.95430076,0.017515946,0.005281585,0.0009502017,0.018333303,0.0036180913],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010139912,0.0025476762,0.0033418885,0.009240819,0.0022587166,0.0093203625,0.0031318406,0.008973027,0.06140101],"category_scores_gemma":[0.08146083,0.0013876284,0.0018406879,0.004655824,0.0024060819,0.0074617486,0.0024756582,0.006956439,0.034273755],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002990342,0.0000073148053,0.00003044974,0.0006895703,0.000011571644,0.00012517953,0.000038140442,0.000011070641,0.00005825954,0.00013530166,0.9905087,0.008354547],"study_design_scores_gemma":[0.00007735338,0.00003664856,0.00022627013,0.0025479656,0.000041285486,0.00049387896,0.00017028954,0.00006600579,0.00015400673,0.00043640067,0.9957241,0.000025729329],"about_ca_topic_score_codex":0.0016490043,"about_ca_topic_score_gemma":0.0033275855,"teacher_disagreement_score":0.98986006,"about_ca_system_score_codex":0.0043037073,"about_ca_system_score_gemma":0.005015376,"threshold_uncertainty_score":0.20540684},"labels":[],"label_agreement":null},{"id":"W4213288981","doi":"10.1093/gigascience/giac003","title":"Future-proofing and maximizing the utility of metadata: The PHA4GE SARS-CoV-2 contextual data specification package","year":2022,"lang":"en","type":"article","venue":"GigaScience","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":45,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Dalhousie University; McMaster University; BC Centre for Disease Control; Simon Fraser University","funders":"Biotechnology and Biological Sciences Research Council; National Institutes of Health; U.S. National Library of Medicine; Wellcome Trust; Bill and Melinda Gates Foundation","keywords":"Interoperability; Computer science; Metadata; Harmonization; Standardization; Consistency (knowledge bases); Data science; Open science; Data sharing; Openness to experience; Best practice; Data integration; World Wide Web; Data mining; Medicine","score_opus":0.10222059033546696,"score_gpt":0.32143076272178844,"score_spread":0.21921017238632148,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213288981","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0058971634,0.0003195245,0.8995104,0.0056904173,0.00050862983,0.0015395617,0.026363846,0.048269514,0.011900936],"genre_scores_gemma":[0.030179838,0.00062507094,0.87306195,0.0025002335,0.00024593933,0.0017196615,0.07592403,0.010895864,0.0048473906],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98000526,0.0074068275,0.004116584,0.0019168082,0.005532141,0.0010223523],"domain_scores_gemma":[0.9469278,0.015178762,0.0034447312,0.02203972,0.010328514,0.0020804366],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.042890448,0.0014820901,0.0009386878,0.004534183,0.0018146693,0.0066753402,0.004109139,0.0024620173,0.007603279],"category_scores_gemma":[0.058108088,0.001534492,0.00276845,0.0035248387,0.00230259,0.008320876,0.009239226,0.0041024163,0.0077819163],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010147492,0.00051937095,0.019686477,0.0019420426,0.0002988145,0.0010818087,0.0036724012,0.015152514,0.01874776,0.35205868,0.34961286,0.2362125],"study_design_scores_gemma":[0.00016917981,0.00015278497,0.0032483917,0.001344132,0.00012934998,0.00058976334,0.0008139704,0.023331178,0.020135593,0.11304964,0.83676964,0.00026637304],"about_ca_topic_score_codex":0.010916441,"about_ca_topic_score_gemma":0.009678124,"teacher_disagreement_score":0.9571096,"about_ca_system_score_codex":0.002625597,"about_ca_system_score_gemma":0.012733698,"threshold_uncertainty_score":0.22682905},"labels":[],"label_agreement":null},{"id":"W4213348372","doi":"10.1038/npre.2011.6043.1","title":"Acknowledging contributions to online expert assistance","year":2011,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"","keywords":"Reputation; Computer science; Quality (philosophy); Judgement; World Wide Web; Closing (real estate); Internet privacy; Web site; Information retrieval; The Internet; Sociology; Political science; Epistemology","score_opus":0.01851932548750148,"score_gpt":0.3454918715416348,"score_spread":0.32697254605413334,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4213348372","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"incentives","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"incentives","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1466184,0.0045990623,0.13152908,0.16530517,0.013884505,0.001404046,0.0044819936,0.016817762,0.51536],"genre_scores_gemma":[0.65444255,0.0020764496,0.09421369,0.025212964,0.0060123955,0.0012118616,0.0032109194,0.0036372931,0.2099819],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9499721,0.03355687,0.0020739106,0.002332331,0.009971189,0.0020937007],"domain_scores_gemma":[0.66101617,0.23844594,0.015622766,0.024567733,0.041280776,0.019066634],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.035187535,0.0008068294,0.00059058657,0.003866464,0.0040285923,0.008476137,0.0021497756,0.0044251666,0.08334959],"category_scores_gemma":[0.19333248,0.00040404464,0.0005870747,0.002329157,0.0025507743,0.007904726,0.012420773,0.0023981838,0.032831505],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005438996,0.00038120087,0.008934558,0.0016387078,0.000062290375,0.0011096552,0.031776596,0.00075878325,0.0030072962,0.02445033,0.5051554,0.4221813],"study_design_scores_gemma":[0.000077939905,0.00020004857,0.0043288013,0.0009917915,0.000033695156,0.0005545412,0.016243724,0.0019918194,0.0024699098,0.01997811,0.95300657,0.00012317514],"about_ca_topic_score_codex":0.0008360553,"about_ca_topic_score_gemma":0.0016046552,"teacher_disagreement_score":0.99152386,"about_ca_system_score_codex":0.002159334,"about_ca_system_score_gemma":0.003738781,"threshold_uncertainty_score":0.27883214},"labels":[],"label_agreement":null},{"id":"W4220795730","doi":"10.3389/ftox.2022.817999","title":"Implementation of Zebrafish Ontologies for Toxicology Screening","year":2022,"lang":"en","type":"article","venue":"Frontiers in Toxicology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"National Institute of Environmental Health Sciences","keywords":"Terminology; Zebrafish; Computer science; Ontology; Consistency (knowledge bases); Data science; Annotation; Ambiguity; Standardization; Computational biology; Biology; Artificial intelligence; Genetics","score_opus":0.019816685291415743,"score_gpt":0.3095915528756957,"score_spread":0.28977486758427995,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220795730","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013440429,0.00029562347,0.891311,0.0023058192,0.00022167667,0.0023402327,0.023273455,0.03916734,0.027644472],"genre_scores_gemma":[0.06281146,0.00062904065,0.88929284,0.00077663374,0.00003139421,0.0015230307,0.036575716,0.0027770833,0.0055828136],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9937157,0.0015604466,0.0014785905,0.00083415787,0.002151944,0.00025913658],"domain_scores_gemma":[0.98624647,0.0035528457,0.0015809102,0.00513654,0.00302638,0.00045683747],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015424834,0.0010655412,0.0005891055,0.0052842028,0.0017384364,0.003391354,0.0029605788,0.0013547487,0.0052980892],"category_scores_gemma":[0.023130968,0.0010231535,0.0021300327,0.0034771184,0.0013073182,0.0059138495,0.004862837,0.0023421114,0.0032018884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035311113,0.0008684873,0.021078655,0.0032790469,0.00048643022,0.0016091904,0.0076811262,0.013756452,0.091082215,0.22763328,0.10205259,0.53011936],"study_design_scores_gemma":[0.000090379806,0.00020172453,0.020823497,0.0016291266,0.00026983724,0.0008528046,0.0022482965,0.05142106,0.059321158,0.103318416,0.75951695,0.0003067659],"about_ca_topic_score_codex":0.017201247,"about_ca_topic_score_gemma":0.027797801,"teacher_disagreement_score":0.017201247,"about_ca_system_score_codex":0.0040497407,"about_ca_system_score_gemma":0.008577344,"threshold_uncertainty_score":0.081575274},"labels":[],"label_agreement":null},{"id":"W4221090405","doi":"10.1186/s12859-022-04636-8","title":"The Xenopus phenotype ontology: bridging model organism phenotype data to human health and development","year":2022,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; National Institute of Child Health and Human Development; National Human Genome Research Institute; National Institutes of Health","keywords":"Ontology; Open Biomedical Ontologies; Computer science; Ontology-based data integration; Interoperability; Process ontology; Phenotype; Computational biology; Biology; Suggested Upper Merged Ontology; Bioinformatics; Information retrieval; World Wide Web; Semantic Web; Genetics; Gene","score_opus":0.07694043601688808,"score_gpt":0.31537138656195013,"score_spread":0.23843095054506205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221090405","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015757648,0.0009967863,0.9139767,0.004509814,0.0003991935,0.00060860306,0.021237474,0.01221425,0.030299557],"genre_scores_gemma":[0.10193569,0.0028601515,0.8472241,0.0015983763,0.0001502135,0.00090886495,0.034670852,0.0032796585,0.0073721376],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984515,0.0004227312,0.00033024565,0.00027665967,0.0004394051,0.00007936932],"domain_scores_gemma":[0.99672544,0.0012243566,0.0004565635,0.00092020055,0.00048343217,0.00018985037],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034935027,0.0006083588,0.00044044942,0.0022123768,0.000856262,0.0022249732,0.0014803882,0.000883883,0.0043044044],"category_scores_gemma":[0.004999762,0.00045144832,0.0013066852,0.0019882803,0.0014857484,0.0040825717,0.002236216,0.0014432542,0.0014974425],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003537546,0.00014927903,0.012665457,0.003121314,0.00020984869,0.0015365989,0.0040573045,0.009229808,0.03166692,0.60239923,0.09586968,0.23874077],"study_design_scores_gemma":[0.000048802645,0.00006999722,0.007621306,0.0011140807,0.00013050827,0.0014319457,0.0009546669,0.014323514,0.017520469,0.10675348,0.84992564,0.000105619896],"about_ca_topic_score_codex":0.009337616,"about_ca_topic_score_gemma":0.010364395,"teacher_disagreement_score":0.009337616,"about_ca_system_score_codex":0.0018122457,"about_ca_system_score_gemma":0.0038495134,"threshold_uncertainty_score":0.01856649},"labels":[],"label_agreement":null},{"id":"W4221110957","doi":"10.1016/j.artmed.2022.102284","title":"Word-level text highlighting of medical texts for telehealth services","year":2022,"lang":"en","type":"article","venue":"Artificial Intelligence in Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":24,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"Mitacs","keywords":"Computer science; Word2vec; Workload; Context (archaeology); Telehealth; Word (group theory); tf–idf; Domain (mathematical analysis); Digitization; Quality (philosophy); Artificial intelligence; Data science; Information retrieval; Health care; Telemedicine; Term (time)","score_opus":0.0774513499730896,"score_gpt":0.37700745367969685,"score_spread":0.29955610370660724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221110957","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4365891,0.006676306,0.32285616,0.007415469,0.0032711609,0.0022661716,0.11789811,0.04832697,0.054700457],"genre_scores_gemma":[0.5468635,0.0019730143,0.38808894,0.0008598443,0.0008193486,0.00046879065,0.04311452,0.0027970886,0.015015],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99959856,0.00009067425,0.000072421266,0.00009173215,0.00011011659,0.00003644193],"domain_scores_gemma":[0.99564123,0.0026562838,0.00052757835,0.00015823034,0.0008166782,0.00019995647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00053654076,0.00078775303,0.00030939048,0.0035148219,0.0005200613,0.0012871418,0.00035936543,0.00070438354,0.016706394],"category_scores_gemma":[0.0047683064,0.00014971558,0.0003769892,0.002189273,0.00025418476,0.0013506764,0.0007734874,0.0006171183,0.0056991884],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022086233,0.00025213975,0.0075237746,0.005695604,0.00011059295,0.0025167959,0.004152983,0.0020783625,0.27492565,0.009711218,0.11258668,0.5782377],"study_design_scores_gemma":[0.00025658592,0.0007366486,0.047462985,0.002129075,0.0005240116,0.0036954419,0.005203775,0.06212589,0.22185116,0.01644805,0.639378,0.00018839446],"about_ca_topic_score_codex":0.00095232914,"about_ca_topic_score_gemma":0.0015745545,"teacher_disagreement_score":0.016706394,"about_ca_system_score_codex":0.00038933932,"about_ca_system_score_gemma":0.00080052746,"threshold_uncertainty_score":0.055888474},"labels":[],"label_agreement":null},{"id":"W4224037213","doi":"10.1101/2022.04.13.22273750","title":"Mondo: Unifying diseases for the world, by the world","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":102,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bioinformatics Solutions (Canada); Jewish General Hospital","funders":"U.S. National Library of Medicine; NIH Office of the Director; National Human Genome Research Institute; National Institutes of Health","keywords":"Ontology; Disease; Computer science; Data science; Key (lock); Data integration; Medicine; Data mining; Computer security","score_opus":0.029588305509629893,"score_gpt":0.312161597043262,"score_spread":0.28257329153363214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224037213","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021319354,0.0069336565,0.7444338,0.013812106,0.002733594,0.00075352675,0.12791872,0.028827637,0.053267606],"genre_scores_gemma":[0.119288124,0.0053709745,0.73099685,0.002209052,0.0006502354,0.0006950279,0.12604924,0.0036775232,0.011062944],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99884796,0.00027569706,0.00015902618,0.0002877163,0.0003438932,0.000085698324],"domain_scores_gemma":[0.99871707,0.00040050145,0.00015755054,0.00038075366,0.00015727314,0.00018695144],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023094907,0.0009288791,0.0006533263,0.005239013,0.0014315882,0.0042075827,0.0011084751,0.0010085726,0.011833126],"category_scores_gemma":[0.00837669,0.00043318924,0.002146503,0.0041474057,0.0009151904,0.0045443643,0.006426963,0.0014860436,0.0038348564],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00036928282,0.00007237049,0.008965746,0.001404937,0.00022232012,0.0007754053,0.0015446469,0.004618558,0.0028628968,0.5237163,0.255147,0.20030047],"study_design_scores_gemma":[0.000051893356,0.000018582228,0.0030049102,0.00048084333,0.00008431224,0.00061830954,0.0005596919,0.015195271,0.001194371,0.20044267,0.77829736,0.00005191059],"about_ca_topic_score_codex":0.006959928,"about_ca_topic_score_gemma":0.008954654,"teacher_disagreement_score":0.011833126,"about_ca_system_score_codex":0.0012356268,"about_ca_system_score_gemma":0.0027042017,"threshold_uncertainty_score":0.03958577},"labels":[],"label_agreement":null},{"id":"W4224044074","doi":"10.1101/2022.03.29.22273096","title":"VentRa. Validation study of the ventricle feature estimation and classification tool to differentiate behavioral variant frontotemporal dementia from psychiatric disorders and other degenerative diseases","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Douglas Mental Health University Institute; McGill University; Montreal Neurological Institute and Hospital","funders":"","keywords":"Frontotemporal dementia; Cohort; Dementia; Medicine; Atrophy; Vascular dementia; Disease; Frontotemporal lobar degeneration; Psychiatry; Psychology; Audiology; Pathology","score_opus":0.022612559243413894,"score_gpt":0.2959614134122825,"score_spread":0.2733488541688686,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224044074","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94113034,0.0016128187,0.021893533,0.00027130215,0.00035124028,0.0008894451,0.027201261,0.0040326198,0.0026174774],"genre_scores_gemma":[0.865199,0.0002872969,0.037549466,0.00033125273,0.0002260206,0.00076251326,0.09279833,0.000511549,0.0023345682],"study_design_codex":"observational","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9960867,0.0013846494,0.000356287,0.001300514,0.0006515071,0.00022031365],"domain_scores_gemma":[0.9932435,0.002447121,0.001033983,0.0012403879,0.0014658767,0.00056914287],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010318034,0.0016445818,0.0006953497,0.0023178877,0.00043332405,0.0012137578,0.002313228,0.0013328126,0.0016199966],"category_scores_gemma":[0.0128977,0.00040567026,0.0012673042,0.0005796346,0.0005694377,0.0007262183,0.001532037,0.00079904473,0.001818336],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009476008,0.0025545799,0.57024884,0.0016212999,0.0038816964,0.0010683995,0.0009767386,0.02255649,0.04014758,0.0014751027,0.0706302,0.2753631],"study_design_scores_gemma":[0.0027860766,0.0059605655,0.6662619,0.00065800623,0.0012220743,0.005199212,0.0007656938,0.24834684,0.035239324,0.0017732152,0.03147394,0.00031319965],"about_ca_topic_score_codex":0.0057993857,"about_ca_topic_score_gemma":0.009195066,"teacher_disagreement_score":0.010318034,"about_ca_system_score_codex":0.00047987528,"about_ca_system_score_gemma":0.00085361564,"threshold_uncertainty_score":0.054567695},"labels":[],"label_agreement":null},{"id":"W4224303302","doi":"10.2196/37804","title":"Conditional Probability Joint Extraction of Nested Biomedical Events: Design of a Unified Extraction Framework Based on Neural Networks","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Natural Science Foundation of China","keywords":"Computer science; Event (particle physics); Artificial intelligence; Conditional random field; Biomedical text mining; Joint probability distribution; Machine learning; Natural language processing; Data mining; Text mining; Mathematics","score_opus":0.03417981275461371,"score_gpt":0.3165629690061541,"score_spread":0.2823831562515404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224303302","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062159295,0.00017569306,0.9898499,0.00020211641,0.000022258486,0.000106324675,0.0002181439,0.002290869,0.0009186562],"genre_scores_gemma":[0.23811829,0.0005102639,0.75211203,0.00040177125,0.00007686007,0.0004223988,0.0023845192,0.000345522,0.0056282734],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99915004,0.00011253645,0.000070987895,0.00039629685,0.00016836978,0.00010166743],"domain_scores_gemma":[0.99910957,0.0003428543,0.00010708375,0.000092131144,0.00029526054,0.000053082746],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001743192,0.0015246542,0.00090796786,0.0024305473,0.00064510736,0.0013297183,0.0023929419,0.0014540767,0.0028805728],"category_scores_gemma":[0.0030536607,0.00074568426,0.0016207981,0.001359318,0.0007718653,0.003210158,0.0018746931,0.0018574302,0.0011081505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002801618,0.00026363603,0.004199593,0.0002828583,0.00025893413,0.0005828508,0.00025649188,0.35085773,0.022616182,0.027595025,0.008165479,0.584641],"study_design_scores_gemma":[0.000005428685,0.000017175349,0.00032360962,0.000012632342,0.000034953606,0.000047835005,0.000014417896,0.9866245,0.003947526,0.0075493967,0.0014118131,0.00001072364],"about_ca_topic_score_codex":0.016823377,"about_ca_topic_score_gemma":0.025322689,"teacher_disagreement_score":0.016823377,"about_ca_system_score_codex":0.0017977469,"about_ca_system_score_gemma":0.0025870698,"threshold_uncertainty_score":0.0334509},"labels":[],"label_agreement":null},{"id":"W4224308858","doi":"10.1145/3485447.3511943","title":"QEN: Applicable Taxonomy Completion via Evaluating Full Taxonomic Relations","year":2022,"lang":"en","type":"article","venue":"Proceedings of the ACM Web Conference 2022","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Computer Research Institute of Montréal","funders":"","keywords":"Taxonomy (biology); Computer science; Artificial intelligence; Information retrieval; Natural language processing; Biology; Zoology","score_opus":0.05760903710980053,"score_gpt":0.27508796583559464,"score_spread":0.2174789287257941,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224308858","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.037507802,0.002333906,0.9322351,0.00076154707,0.00026621815,0.00047753897,0.005136844,0.01546317,0.005817799],"genre_scores_gemma":[0.25097647,0.0011111781,0.711478,0.00071455305,0.0001305803,0.000390249,0.022150816,0.0009877444,0.012060362],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982911,0.0003146774,0.00010910519,0.00065448007,0.0004834217,0.00014727382],"domain_scores_gemma":[0.9972555,0.0011219755,0.00021346907,0.0005245407,0.00073685084,0.00014766685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023066269,0.0018877975,0.0011868626,0.0031549465,0.00081332435,0.00147635,0.0035590813,0.0017834736,0.0069635026],"category_scores_gemma":[0.01094336,0.0006391494,0.0012148346,0.0022343018,0.0006975915,0.0066517093,0.0027671761,0.002262231,0.00316407],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044412495,0.0003536816,0.0060079047,0.00092173525,0.00016809242,0.00038910806,0.0004461282,0.056542937,0.0122560365,0.01248368,0.04581424,0.86417234],"study_design_scores_gemma":[0.000042656993,0.00015273594,0.0019385818,0.00010878892,0.0000706759,0.00030596333,0.00018230153,0.9326392,0.008307786,0.037043605,0.019155573,0.000052150725],"about_ca_topic_score_codex":0.018217057,"about_ca_topic_score_gemma":0.032840706,"teacher_disagreement_score":0.018217057,"about_ca_system_score_codex":0.0018113505,"about_ca_system_score_gemma":0.0020334704,"threshold_uncertainty_score":0.03622204},"labels":[],"label_agreement":null},{"id":"W4224646058","doi":"10.18653/v1/2022.bionlp-1.32","title":"ICDBigBird: A Contextual Embedding Model for ICD Code Classification","year":2022,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; University of Toronto","funders":"Mitacs","keywords":"Computer science; Word embedding; Task (project management); Embedding; Process (computing); Artificial intelligence; Health care; ICD-10; Machine learning; Graph; Natural language processing; Theoretical computer science","score_opus":0.06264871263909677,"score_gpt":0.3368751190488523,"score_spread":0.27422640640975554,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4224646058","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015103939,0.0019550368,0.92461324,0.0013558664,0.0005653553,0.00048510596,0.033551987,0.018217154,0.0041523594],"genre_scores_gemma":[0.16571875,0.0015324327,0.7414713,0.0011382563,0.00022478275,0.0009759378,0.08032516,0.0015155792,0.007097811],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99899966,0.00030830686,0.00009117975,0.0002805659,0.0002455918,0.00007469071],"domain_scores_gemma":[0.9988482,0.0004890296,0.00007133034,0.00026275014,0.00026816837,0.000060555816],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013635338,0.00095134473,0.00077693386,0.0018471543,0.0006374038,0.001378165,0.0016133392,0.0010187291,0.0062295375],"category_scores_gemma":[0.006033785,0.00038777007,0.0014455933,0.001737892,0.0003861735,0.0020697962,0.0024631172,0.00190284,0.003569027],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078164093,0.0004258118,0.009938711,0.0010502923,0.00049224397,0.00043786495,0.00054265355,0.06114379,0.004050682,0.027263112,0.13804905,0.7558241],"study_design_scores_gemma":[0.000129259,0.00018097421,0.0038629062,0.0003067422,0.00021420135,0.00039328696,0.00031390006,0.78424686,0.005366744,0.09076396,0.11411221,0.00010892902],"about_ca_topic_score_codex":0.015794631,"about_ca_topic_score_gemma":0.02982164,"teacher_disagreement_score":0.015794631,"about_ca_system_score_codex":0.0009020293,"about_ca_system_score_gemma":0.0015259172,"threshold_uncertainty_score":0.03140539},"labels":[],"label_agreement":null},{"id":"W4225003357","doi":"10.2196/35789","title":"Generation of a Fast Healthcare Interoperability Resources (FHIR)-based Ontology for Federated Feasibility Queries in the Context of COVID-19: Feasibility Study","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":33,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Interoperability; Ontology; Usability; Terminology; Context (archaeology); Information retrieval; Database; Data mining; World Wide Web","score_opus":0.10320209411290962,"score_gpt":0.3926039160610196,"score_spread":0.2894018219481099,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225003357","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056965064,0.00014349956,0.8827308,0.0015597997,0.00018429224,0.0017703127,0.014180342,0.02945248,0.013013537],"genre_scores_gemma":[0.15495832,0.00022418595,0.8007308,0.0004018135,0.00003464845,0.0010511087,0.032184567,0.005723037,0.0046914867],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966012,0.00090838317,0.0005454615,0.0005081247,0.0011468778,0.00028981853],"domain_scores_gemma":[0.99439853,0.0025518006,0.0002739047,0.00093692774,0.0015836091,0.00025517194],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0052052382,0.00093455473,0.0006513191,0.0026325502,0.0011059211,0.0034639305,0.0015095457,0.0012611096,0.0050047343],"category_scores_gemma":[0.011923456,0.00063496275,0.0026592643,0.0016263601,0.0008529034,0.004597144,0.0031726267,0.001730657,0.0019628138],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018685252,0.0017870506,0.019309709,0.003995664,0.00037894546,0.00643449,0.019396957,0.072149344,0.13370678,0.25845033,0.119952455,0.36256975],"study_design_scores_gemma":[0.00027816131,0.0003330524,0.005482642,0.0007037344,0.00023822617,0.0019896592,0.0045404085,0.44004345,0.10655827,0.057595566,0.38187644,0.00036043295],"about_ca_topic_score_codex":0.011740904,"about_ca_topic_score_gemma":0.008060463,"teacher_disagreement_score":0.011740904,"about_ca_system_score_codex":0.0025017895,"about_ca_system_score_gemma":0.0038954995,"threshold_uncertainty_score":0.027528286},"labels":[],"label_agreement":null},{"id":"W4225286402","doi":"10.21203/rs.3.rs-1465079/v1","title":"Features of a FAIR vocabulary","year":2022,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"Engineering and Physical Sciences Research Council; European Commission; EOSC-Life; European Molecular Biology Laboratory","keywords":"Vocabulary; Computer science; Controlled vocabulary; Fair use; Information retrieval; Artificial intelligence; Natural language processing; Linguistics","score_opus":0.056363696896941075,"score_gpt":0.4119136454336711,"score_spread":0.35554994853673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225286402","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.084485,0.00084307767,0.83903885,0.006786508,0.00028672072,0.0014020775,0.0022090499,0.0015740297,0.06337466],"genre_scores_gemma":[0.66207916,0.0002537569,0.3279557,0.00091597013,0.00014926484,0.001234603,0.0022117447,0.00035433721,0.004845415],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9614372,0.013549568,0.0065174713,0.003953723,0.012733226,0.0018089071],"domain_scores_gemma":[0.8960821,0.04701124,0.009343602,0.021309821,0.023343392,0.0029099472],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.034499772,0.00081332674,0.0010393022,0.008024396,0.004180719,0.009679961,0.0020992204,0.0021827943,0.0059037297],"category_scores_gemma":[0.13359077,0.0006464901,0.0018125053,0.004483757,0.010146127,0.021218564,0.009446479,0.0030481233,0.0008986411],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001269022,0.00007910821,0.009816595,0.00030958504,0.00006762962,0.0001531343,0.0028388754,0.0046239407,0.001787463,0.90623504,0.005489882,0.06847189],"study_design_scores_gemma":[0.00003194327,0.00010000819,0.0035369762,0.000524038,0.00007164761,0.00023816805,0.001867113,0.015436772,0.0036607392,0.9315339,0.04291977,0.00007894515],"about_ca_topic_score_codex":0.009943457,"about_ca_topic_score_gemma":0.006692245,"teacher_disagreement_score":0.9979008,"about_ca_system_score_codex":0.005139065,"about_ca_system_score_gemma":0.007514837,"threshold_uncertainty_score":0.18245447},"labels":[],"label_agreement":null},{"id":"W4225373816","doi":"10.32473/flairs.v35i.130660","title":"Protein-Protein Interaction Extraction using Attention-based Tree-Structured Neural Network Models","year":2022,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Tree (set theory); Artificial intelligence; Task (project management); Artificial neural network; Machine learning; Natural language processing; Phrase; Recurrent neural network; Tree structure; Data structure; Mathematics","score_opus":0.15386596719488524,"score_gpt":0.38203420285943096,"score_spread":0.22816823566454572,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4225373816","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17967996,0.002129759,0.8086079,0.00084761693,0.0001106896,0.00013802001,0.0011075122,0.0028874062,0.0044912477],"genre_scores_gemma":[0.8353425,0.0008057805,0.15500788,0.00034704257,0.00009313977,0.00016012583,0.0025015192,0.00009752906,0.005644571],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999731,0.000054949207,0.00001811034,0.00009958322,0.000057306108,0.000038994487],"domain_scores_gemma":[0.9991033,0.0005606861,0.00009700721,0.000039475497,0.00016496477,0.00003468941],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00060774625,0.000828867,0.00084900163,0.0015104242,0.0003412846,0.00073926325,0.001068941,0.0010719488,0.0014741732],"category_scores_gemma":[0.0019840205,0.00036222974,0.001110589,0.0015938788,0.00026479823,0.0017589243,0.00063515676,0.0009684851,0.0006564757],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043459726,0.00041100328,0.0048666107,0.00018742365,0.00020322343,0.00041814905,0.00014369817,0.68104845,0.00900312,0.0077878325,0.0063702133,0.28912574],"study_design_scores_gemma":[0.0000034468965,0.000010831366,0.00017303399,0.0000025690117,0.000008670473,0.000010410364,0.0000034505276,0.9969292,0.00027505023,0.002450881,0.0001301398,0.000002259573],"about_ca_topic_score_codex":0.014504697,"about_ca_topic_score_gemma":0.018763158,"teacher_disagreement_score":0.014504697,"about_ca_system_score_codex":0.0012436665,"about_ca_system_score_gemma":0.001098144,"threshold_uncertainty_score":0.028840542},"labels":[],"label_agreement":null},{"id":"W4226100175","doi":"10.1177/14604582221083850","title":"Developing a pneumonia diagnosis ontology from multiple knowledge sources","year":2022,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Université du Québec en Outaouais","funders":"","keywords":"Ontology; Pneumonia; Protégé; Medical diagnosis; Medicine; Open Biomedical Ontologies; Knowledge representation and reasoning; Intensive care medicine; Computer science; Data science; Knowledge management; Domain knowledge; Pathology; Artificial intelligence; Process ontology; Semantic Web; Ontology alignment; Internal medicine","score_opus":0.04636128712420141,"score_gpt":0.3256436618335007,"score_spread":0.2792823747092993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226100175","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040826883,0.001867642,0.87982935,0.0043966104,0.00031764948,0.0033108296,0.041210838,0.009129641,0.019110497],"genre_scores_gemma":[0.06968649,0.0016177624,0.8718386,0.00071284565,0.000053298532,0.0011192625,0.05260482,0.00034005515,0.002026839],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.996393,0.00058217073,0.00096152304,0.00062458724,0.0012638331,0.00017484771],"domain_scores_gemma":[0.99294865,0.0034123363,0.00063117716,0.0009814977,0.0017190243,0.00030724882],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0053330907,0.0009931589,0.00079916086,0.014338337,0.0015692606,0.0035648441,0.0014256835,0.0013620541,0.0022675325],"category_scores_gemma":[0.015933704,0.00087108754,0.0032832436,0.00794302,0.0006275405,0.006098559,0.005632396,0.0018133586,0.00071925606],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038554866,0.0007810343,0.034119774,0.0088176625,0.001568303,0.006686619,0.0051736305,0.021597827,0.021740977,0.096590385,0.053317327,0.7492209],"study_design_scores_gemma":[0.00028640183,0.00019511639,0.024485324,0.0062448913,0.0018719645,0.004019146,0.0049617286,0.16319573,0.028400112,0.1240561,0.64191365,0.00036986667],"about_ca_topic_score_codex":0.0199561,"about_ca_topic_score_gemma":0.032096077,"teacher_disagreement_score":0.0199561,"about_ca_system_score_codex":0.0026789578,"about_ca_system_score_gemma":0.008884299,"threshold_uncertainty_score":0.039679885},"labels":[],"label_agreement":null},{"id":"W4226329328","doi":"10.1093/bioinformatics/btac195","title":"ELIXIR biovalidator for semantic validation of life science metadata","year":2022,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"Vlaamse regering; Fonds Wetenschappelijk Onderzoek; European Bioinformatics Institute","keywords":"Validator; Computer science; License; Schema (genetic algorithms); Elixir (programming language); Information retrieval; Metadata; JSON; Semantic integration; World Wide Web; Programming language; Semantic Web","score_opus":0.03119283094040965,"score_gpt":0.28790835011622395,"score_spread":0.2567155191758143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226329328","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007914277,0.0010484258,0.40264416,0.0013940595,0.0005767332,0.0009628487,0.058294646,0.5065172,0.020647744],"genre_scores_gemma":[0.061313387,0.00096941664,0.49450088,0.0022127957,0.00020579435,0.0021967965,0.2972729,0.12799862,0.013329467],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98710907,0.0029283701,0.0019172933,0.0021690046,0.005325114,0.0005511837],"domain_scores_gemma":[0.9682663,0.011029374,0.0024609189,0.010402225,0.007197386,0.0006437761],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024666661,0.0025089383,0.0014377608,0.006342515,0.0025802942,0.0074865627,0.004270434,0.0024215402,0.03149138],"category_scores_gemma":[0.05426424,0.0016668623,0.0034140702,0.0032895505,0.0017419046,0.008659867,0.011805419,0.0038000944,0.030898647],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031235563,0.00038725793,0.013917451,0.005561309,0.0007651137,0.0012739964,0.0028460408,0.0066527063,0.02653859,0.08968955,0.59228015,0.2569643],"study_design_scores_gemma":[0.00030377918,0.00016647647,0.006507653,0.0023086017,0.000249515,0.0009942097,0.00052408804,0.046883408,0.07674799,0.04975083,0.8151497,0.00041384457],"about_ca_topic_score_codex":0.006904767,"about_ca_topic_score_gemma":0.0067571593,"teacher_disagreement_score":0.03149138,"about_ca_system_score_codex":0.0028787956,"about_ca_system_score_gemma":0.0070295907,"threshold_uncertainty_score":0.13045138},"labels":[],"label_agreement":null},{"id":"W4226512137","doi":"10.1504/ijiids.2022.120143","title":"Supporting user-centred ontology visualisation: predictive analytics using eye gaze to enhance human-ontology interaction","year":2022,"lang":"en","type":"article","venue":"International Journal of Intelligent Information and Database Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria; University of British Columbia","funders":"","keywords":"Computer science; Ontology; Visual analytics; Human–computer interaction; Gaze; Visualization; Analytics; Eye tracking; Data science; Artificial intelligence","score_opus":0.038203926064547214,"score_gpt":0.38985262046015434,"score_spread":0.35164869439560714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4226512137","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6692618,0.001067395,0.31213167,0.00055054406,0.00009705659,0.0003517764,0.0019250506,0.0061545186,0.008460124],"genre_scores_gemma":[0.93425006,0.00032488274,0.06310967,0.00005110535,0.00002999176,0.00011580505,0.00051015656,0.00018566386,0.0014226215],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9994116,0.00023050264,0.00002350741,0.00011649877,0.00016136766,0.000056605255],"domain_scores_gemma":[0.99575704,0.0030460092,0.0003491942,0.00030125596,0.00041396084,0.00013248848],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010361894,0.0006888598,0.00046562968,0.0019175065,0.00033878288,0.0015592262,0.0004777111,0.00060940604,0.0032745432],"category_scores_gemma":[0.009832103,0.00022031438,0.0003115765,0.0008593873,0.00026060725,0.0013692075,0.0014561559,0.0005882473,0.00093802484],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027315845,0.0005549823,0.086868644,0.0017199512,0.00026564108,0.0005835571,0.014450208,0.020592535,0.20169494,0.0032911983,0.009213824,0.65803295],"study_design_scores_gemma":[0.0001475775,0.0017672329,0.2904813,0.0005325937,0.0002812329,0.0014812107,0.006847364,0.5726758,0.08753023,0.017150957,0.020630054,0.00047440798],"about_ca_topic_score_codex":0.004434555,"about_ca_topic_score_gemma":0.0061977142,"teacher_disagreement_score":0.004434555,"about_ca_system_score_codex":0.00036327116,"about_ca_system_score_gemma":0.00042401956,"threshold_uncertainty_score":0.01095438},"labels":[],"label_agreement":null},{"id":"W4229838926","doi":"10.4018/978-1-7998-1204-3.ch029","title":"Semantic Reconciliation of Electronic Health Records Using Semantic Web Technologies","year":2019,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"RDF; Semantic reasoner; Computer science; Semantic Web; Information retrieval; Coding (social sciences); Electronic medical record; Semantics (computer science); SPARQL; Health records; Medical record; Semantic analytics; Process (computing); Electronic health record; World Wide Web; Semantic Web Stack; Health care; Artificial intelligence; Medicine; Programming language; Internet privacy; Mathematics","score_opus":0.020640231462399693,"score_gpt":0.27699949666546586,"score_spread":0.2563592652030662,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229838926","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007971571,0.0035851686,0.9598382,0.004379134,0.00069174374,0.00041313452,0.0012087886,0.0035734365,0.018338805],"genre_scores_gemma":[0.09516276,0.0043475335,0.8789056,0.0018089433,0.00045743247,0.00030141664,0.007983819,0.0009349776,0.010097566],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9863854,0.005968408,0.0013098024,0.0014840752,0.004488217,0.00036409486],"domain_scores_gemma":[0.9869774,0.0053378004,0.000981075,0.0049094185,0.0016444138,0.00015005651],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012821448,0.0010061216,0.0013515332,0.010228568,0.0018928843,0.008536509,0.003752325,0.0024360002,0.0038614667],"category_scores_gemma":[0.016670018,0.00064734335,0.0028531174,0.01192162,0.0026841355,0.015304004,0.007451503,0.0030000012,0.0018407082],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015960465,0.00020706296,0.0014859902,0.0014262003,0.00037759024,0.0011086267,0.003993581,0.0107417535,0.0043040337,0.39881596,0.02525419,0.55212533],"study_design_scores_gemma":[0.000068803274,0.00007435316,0.0013366444,0.001548178,0.00029710022,0.0014760129,0.002837145,0.052684706,0.018829389,0.47634518,0.44437,0.00013254756],"about_ca_topic_score_codex":0.0020976278,"about_ca_topic_score_gemma":0.0023551905,"teacher_disagreement_score":0.012821448,"about_ca_system_score_codex":0.0019610594,"about_ca_system_score_gemma":0.0039995373,"threshold_uncertainty_score":0.06780714},"labels":[],"label_agreement":null},{"id":"W4229847747","doi":"10.29173/irie196","title":"On IRIE Vol. 5","year":2006,"lang":"en","type":"article","venue":"The International Review of Information Ethics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Geography","score_opus":0.020643678237918083,"score_gpt":0.32411059613919163,"score_spread":0.30346691790127356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4229847747","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00067725824,0.0050557875,0.0002905081,0.010492939,0.0118751945,0.00012948082,0.0008538923,0.00035047746,0.9702745],"genre_scores_gemma":[0.0019262789,0.0021854625,0.00018304824,0.008145869,0.0019982024,0.000053109725,0.00041898378,0.00010538405,0.9849836],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99888617,0.00014460416,0.00006241249,0.00021273605,0.00048989826,0.00020422213],"domain_scores_gemma":[0.99902976,0.0001949353,0.000049978615,0.00017755275,0.00037020983,0.0001774466],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0013567952,0.0010149566,0.00083327445,0.0022892763,0.0027073633,0.008647558,0.0012479157,0.004186335,0.62310284],"category_scores_gemma":[0.0040700296,0.00044893933,0.0009200234,0.0014778945,0.0016140307,0.002309628,0.0032816397,0.0048040156,0.5034107],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007044089,0.000088077395,0.00046594144,0.00021831093,0.000012334126,0.00021323585,0.00011323143,0.00006839506,0.0005930784,0.01264538,0.93015784,0.055353764],"study_design_scores_gemma":[0.000004066881,0.000009948398,0.00040533548,0.00014445357,0.0000021764708,0.000046526704,0.000047719666,0.000017080574,0.00005853738,0.0005247004,0.9987361,0.0000033509193],"about_ca_topic_score_codex":0.005314627,"about_ca_topic_score_gemma":0.012377481,"teacher_disagreement_score":0.62310284,"about_ca_system_score_codex":0.0032636668,"about_ca_system_score_gemma":0.0020733005,"threshold_uncertainty_score":0.537598},"labels":[],"label_agreement":null},{"id":"W4230225705","doi":"10.1038/npre.2010.5443","title":"Keynote: A renaissance for the point mutation: from legacy data to semantic web service","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Computer science; Mutation; Information retrieval; Annotation; Population; World Wide Web; Visualization; Data mining; Artificial intelligence; Genetics; Biology","score_opus":0.03738788146197312,"score_gpt":0.33241611869051185,"score_spread":0.2950282372285387,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230225705","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005205094,0.0070476173,0.54134536,0.3074764,0.08154525,0.00027842223,0.002293287,0.010431474,0.044377074],"genre_scores_gemma":[0.11848575,0.021973932,0.41114777,0.107384145,0.053560805,0.0008494195,0.007986157,0.014370479,0.26424146],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99243355,0.0019392867,0.0005037049,0.0012209892,0.0033037052,0.0005986503],"domain_scores_gemma":[0.98492545,0.0046989415,0.0004556666,0.0034440453,0.0043253703,0.0021505638],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015669435,0.0009651685,0.0009577392,0.0021816636,0.002595217,0.011904174,0.0028514955,0.005780792,0.02333746],"category_scores_gemma":[0.024113417,0.00072218676,0.0016660183,0.0026688601,0.005022889,0.024469577,0.0090646865,0.012525489,0.015309502],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039665287,0.00007737605,0.0007705999,0.0004775198,0.000057959503,0.000583211,0.0015318963,0.0011472246,0.006142481,0.38911685,0.42864427,0.17105389],"study_design_scores_gemma":[0.000023006398,0.000051658662,0.00025267646,0.0002206394,0.000020082407,0.0002991387,0.00045736518,0.0019991295,0.0025857564,0.08065374,0.913375,0.00006166521],"about_ca_topic_score_codex":0.0034160807,"about_ca_topic_score_gemma":0.0025292062,"teacher_disagreement_score":0.02333746,"about_ca_system_score_codex":0.0032315396,"about_ca_system_score_gemma":0.0038623377,"threshold_uncertainty_score":0.082868874},"labels":[],"label_agreement":null},{"id":"W4230902601","doi":"10.18733/cpi29490","title":"Contributor Bioggraphies","year":2019,"lang":"en","type":"article","venue":"Cultural and Pedagogical Inquiry","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"History; Geography","score_opus":0.39791319814539355,"score_gpt":0.43863406014322015,"score_spread":0.040720861997826596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230902601","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014220254,0.0031847816,0.008653902,0.032509457,0.06507341,0.0014253891,0.26425505,0.007895404,0.6155806],"genre_scores_gemma":[0.011743176,0.004352739,0.010476275,0.005289224,0.010958606,0.0020156356,0.15034163,0.006789406,0.7980332],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9966156,0.00055916415,0.00029541412,0.000609844,0.0016810654,0.00023886],"domain_scores_gemma":[0.9696979,0.005701271,0.001496956,0.0035320078,0.01739281,0.0021789868],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0029655078,0.0012751841,0.0010809586,0.02172487,0.0037790828,0.0062297178,0.002232384,0.00218439,0.6366888],"category_scores_gemma":[0.054142956,0.00071645394,0.0007183936,0.021179227,0.00075942127,0.005452245,0.005178209,0.0028395406,0.38122877],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007985337,0.000005901466,0.00008302459,0.00010980055,0.0000012594398,0.000020496793,0.00006375012,0.000015482194,0.000038243266,0.00092802546,0.98568785,0.013038107],"study_design_scores_gemma":[0.0000044922035,0.000002157184,0.0002084368,0.00013602452,0.0000031150905,0.00003286643,0.00009073431,0.000052777228,0.000060923685,0.0007780923,0.99862516,0.0000052052305],"about_ca_topic_score_codex":0.0145833455,"about_ca_topic_score_gemma":0.017255323,"teacher_disagreement_score":0.36331117,"about_ca_system_score_codex":0.003403607,"about_ca_system_score_gemma":0.0079314895,"threshold_uncertainty_score":0.51821923},"labels":[],"label_agreement":null},{"id":"W4231195059","doi":"10.1345/aph.1h463b","title":"Reply: Community Identification of Natural Health Product-Drug Interactions","year":2007,"lang":"en","type":"article","venue":"Annals of Pharmacotherapy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Health Canada; University of Alberta","funders":"","keywords":"Medicine; Natural product; Identification (biology); Drug; Pharmacology; Stereochemistry","score_opus":0.0742980701981207,"score_gpt":0.4462531306522291,"score_spread":0.3719550604541084,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231195059","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006088525,0.0007312358,0.00041265617,0.9823722,0.015188152,0.000013865185,0.00031391086,0.000030162913,0.0003287926],"genre_scores_gemma":[0.006945515,0.0012405469,0.0013756495,0.9652942,0.021031054,0.00008280032,0.0002979287,0.000026160391,0.0037062853],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968997,0.0009966339,0.00058597326,0.00049237726,0.0006465548,0.0003786513],"domain_scores_gemma":[0.96948874,0.016338121,0.0012145049,0.0009916227,0.009727773,0.0022391828],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061859298,0.0007112219,0.0011518816,0.001658276,0.002654275,0.002551518,0.0016447831,0.024759144,0.008268661],"category_scores_gemma":[0.039708838,0.0006674217,0.0013594435,0.0013722178,0.0023038117,0.003933307,0.0028879931,0.024652164,0.005224431],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013385256,0.00003871518,0.0016794743,0.00020214295,0.000077818164,0.00051044195,0.0002886086,0.000082428705,0.00043243446,0.0011391783,0.98555344,0.009861437],"study_design_scores_gemma":[0.00033466762,0.00010251094,0.0079592075,0.0004861859,0.00030383014,0.0026756884,0.0024113369,0.0011234191,0.0008683628,0.013983265,0.9696023,0.00014921713],"about_ca_topic_score_codex":0.004308583,"about_ca_topic_score_gemma":0.0059823515,"teacher_disagreement_score":0.024759144,"about_ca_system_score_codex":0.0016226665,"about_ca_system_score_gemma":0.0033830034,"threshold_uncertainty_score":0.032714725},"labels":[],"label_agreement":null},{"id":"W4231223228","doi":"10.32920/14638710","title":"The trainees' perspective on developing an end-of-grant knowledge translation plan","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Manitoba; Queen's University; McMaster University; University of Toronto; Hamilton Health Sciences; Toronto Metropolitan University; University of Calgary","funders":"Canadian Institutes of Health Research; Fondation pour la Recherche Médicale; European Observatory on Health Systems and Policies","keywords":"Knowledge translation; Craft; Plan (archaeology); Perspective (graphical); Process (computing); Library science; Political science; Medicine; Computer science; Knowledge management; Geography","score_opus":0.09461686856878863,"score_gpt":0.35244418726091675,"score_spread":0.2578273186921281,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231223228","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049777334,0.0006166805,0.014967612,0.951688,0.0026234288,0.00051941857,0.000076438475,0.00011962684,0.024411136],"genre_scores_gemma":[0.2940301,0.0037886512,0.13450748,0.5114185,0.005016836,0.0041628336,0.00039611477,0.0004510788,0.046228427],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.7107527,0.22442833,0.011174103,0.005909307,0.022044573,0.025690883],"domain_scores_gemma":[0.5971288,0.2111573,0.012722928,0.015974686,0.061932746,0.10108344],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.30587605,0.0009281844,0.0009632165,0.0019825592,0.01779482,0.032972913,0.009641612,0.03614795,0.017724415],"category_scores_gemma":[0.30853927,0.0012246788,0.00218483,0.0024535463,0.024448995,0.02313293,0.03273581,0.049379785,0.0053578657],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003118661,0.0010222035,0.0033233643,0.0014239794,0.0000938253,0.0029669649,0.11081159,0.003876629,0.0017411556,0.5137919,0.2549073,0.10572925],"study_design_scores_gemma":[0.00022523958,0.00062293274,0.0013850009,0.0027131448,0.000066863475,0.0008821263,0.0996865,0.003432832,0.0015744589,0.15351138,0.73562044,0.00027907462],"about_ca_topic_score_codex":0.014480207,"about_ca_topic_score_gemma":0.013733644,"teacher_disagreement_score":0.694124,"about_ca_system_score_codex":0.021511449,"about_ca_system_score_gemma":0.19456777,"threshold_uncertainty_score":0.8559784},"labels":[],"label_agreement":null},{"id":"W4231341663","doi":"10.1007/978-3-642-34713-9","title":"Machine Learning and Interpretation in Neuroimaging","year":2012,"lang":"en","type":"book","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":25,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Neuroimaging; Interpretation (philosophy); Artificial intelligence; Cognitive science; Computer science; Psychology; Neuroscience","score_opus":0.009650860820995562,"score_gpt":0.2640385200137727,"score_spread":0.2543876591927771,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231341663","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0051730396,0.069973394,0.8584056,0.014868951,0.0014084777,0.00006752386,0.0006019652,0.0011162183,0.048384808],"genre_scores_gemma":[0.19052796,0.06343086,0.6867297,0.0018951007,0.0029138562,0.00029076755,0.0009642209,0.0005211264,0.052726362],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9993831,0.0003077171,0.000058651935,0.00007615421,0.00015491224,0.000019563915],"domain_scores_gemma":[0.99729925,0.0022057057,0.00008685897,0.00023956632,0.00013705136,0.00003146934],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018078663,0.00048109383,0.0006212628,0.0012195103,0.00025079792,0.002071021,0.0008066757,0.0010327101,0.0065779923],"category_scores_gemma":[0.0066390056,0.00044508136,0.00048122866,0.0014945961,0.0022078499,0.0032493346,0.00088684325,0.0013609476,0.00133904],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022179649,0.000023252409,0.00052266533,0.000622203,0.000046781788,0.00028882467,0.00038097036,0.0040321387,0.0009079526,0.41377893,0.04362847,0.53574556],"study_design_scores_gemma":[0.0000030659398,0.000007020911,0.00053515687,0.00013477053,0.00001060492,0.00039391973,0.000096161035,0.009663339,0.00068422995,0.94861454,0.039845806,0.000011454357],"about_ca_topic_score_codex":0.001152246,"about_ca_topic_score_gemma":0.0018403499,"teacher_disagreement_score":0.0065779923,"about_ca_system_score_codex":0.00064106565,"about_ca_system_score_gemma":0.00079549034,"threshold_uncertainty_score":0.022005618},"labels":[],"label_agreement":null},{"id":"W4231530258","doi":"10.3109/13561820.2011.589227","title":"Canadian Interprofessional Health Collaborative Blog http://cihcblog.com","year":2011,"lang":"en","type":"article","venue":"Journal of Interprofessional Care","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Interprofessional education; Library science; Medical education; World Wide Web; Medicine; Health care; Computer science; Political science","score_opus":0.016646822746495066,"score_gpt":0.311481158238375,"score_spread":0.29483433549188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4231530258","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006446746,0.008665346,0.0017630726,0.08418067,0.009523122,0.00027029085,0.061130587,0.0040481407,0.8239721],"genre_scores_gemma":[0.035176434,0.0068735164,0.0045445254,0.010519233,0.0018120972,0.000096670854,0.021812119,0.0008684328,0.91829705],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989343,0.00004557683,0.000023648423,0.00005609073,0.0007454091,0.00019488564],"domain_scores_gemma":[0.99421906,0.0009684737,0.00018915531,0.00023910786,0.0015667153,0.0028175283],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00094431854,0.0007346586,0.0004620311,0.0057008374,0.005515412,0.005736243,0.0010690428,0.0017721842,0.24780954],"category_scores_gemma":[0.0046980153,0.00027188618,0.00030024027,0.0118826125,0.0014133555,0.0021379406,0.0022361225,0.0018568397,0.035436455],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002162708,0.00001458471,0.00024226111,0.000063581705,0.0000018812602,0.000037725993,0.00011798186,0.000025980204,0.00008181964,0.0016678286,0.97298074,0.024743948],"study_design_scores_gemma":[0.0000118950065,0.0000038807084,0.0011885941,0.000055299326,0.0000037750353,0.00002725631,0.00064432184,0.00007333548,0.00012363182,0.00058687857,0.9972704,0.000010628898],"about_ca_topic_score_codex":0.8057285,"about_ca_topic_score_gemma":0.936848,"teacher_disagreement_score":0.9837702,"about_ca_system_score_codex":0.016229808,"about_ca_system_score_gemma":0.045771718,"threshold_uncertainty_score":0.8290055},"labels":[],"label_agreement":null},{"id":"W4232448836","doi":"10.1515/iupac.87.0058","title":"Aphonia","year":2016,"lang":"en","type":"dataset","venue":"IUPAC Standards Online","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada Research Chairs; University of Toronto","funders":"","keywords":"Glossary; Chemical nomenclature; Relation (database); Psychology; Computer science; Chemistry; Linguistics; Philosophy; Data mining; Organic chemistry","score_opus":0.012187474985428359,"score_gpt":0.3949302460189032,"score_spread":0.3827427710334748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232448836","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00020964978,0.00013826208,0.00021583251,0.00011669937,0.000063381456,0.00003848102,0.99457574,0.0011846519,0.0034572694],"genre_scores_gemma":[0.0004591133,0.00011177385,0.00058272487,0.00013008997,0.000012971627,0.00009194766,0.99615794,0.00013618005,0.0023172277],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99891484,0.00018448912,0.00021127288,0.00030676092,0.0002568065,0.00012578619],"domain_scores_gemma":[0.99788326,0.00055868085,0.00021880471,0.0005891838,0.000541235,0.000208811],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009668528,0.0014502745,0.0010598247,0.003242389,0.0008033645,0.0027049966,0.0018141476,0.0013721667,0.16075677],"category_scores_gemma":[0.0070435903,0.00047861584,0.0011159834,0.004422966,0.0003822183,0.001995138,0.0024021138,0.0014269432,0.22377698],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007821011,0.000016369126,0.0004999976,0.00076115306,0.00001578375,0.000033978806,0.000027680218,0.00011929953,0.00016049358,0.0006223309,0.9909314,0.00673337],"study_design_scores_gemma":[0.000075441945,0.000012309701,0.0016203906,0.00026274883,0.000012554031,0.00008181621,0.000056139237,0.00028585762,0.0003041607,0.0011789997,0.99609417,0.000015304204],"about_ca_topic_score_codex":0.010600603,"about_ca_topic_score_gemma":0.024124471,"teacher_disagreement_score":0.16075677,"about_ca_system_score_codex":0.0011929298,"about_ca_system_score_gemma":0.0021831081,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4232849376","doi":"10.1385/1-59259-038-1:729","title":"DNase I Footprinting","year":2003,"lang":"en","type":"book-chapter","venue":"Humana Press eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hôtel-Dieu de Québec","funders":"","keywords":"Footprinting; DNA footprinting; Deoxyribonuclease I; Hypersensitive site; DNA; Molecular biology; Biology; DNase I hypersensitive site; HMG-box; DNA binding site; Transcription factor; DNA-binding protein; Promoter; Biochemistry; Gene; Gene expression; Base sequence","score_opus":0.06355862030086441,"score_gpt":0.2731416649700737,"score_spread":0.2095830446692093,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232849376","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016483583,0.10856863,0.7683935,0.0015364878,0.0040580384,0.0027006555,0.014346289,0.009961218,0.073951654],"genre_scores_gemma":[0.06352507,0.07555443,0.66938645,0.0021872816,0.00064676366,0.003932013,0.030953407,0.0020936804,0.1517209],"study_design_codex":"bench_or_experimental","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978409,0.00046621118,0.00013529512,0.0007884248,0.00063342863,0.00013580627],"domain_scores_gemma":[0.99919254,0.00036560543,0.00006673179,0.00016299759,0.00016391621,0.000048238777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001324333,0.0020487313,0.0026897246,0.0027894462,0.0010665277,0.0012626078,0.0019370634,0.0013329573,0.031576775],"category_scores_gemma":[0.0011906753,0.001097767,0.0010347915,0.0023905674,0.00094801374,0.0008726373,0.0013714097,0.0033376715,0.032269012],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003067268,0.00015549784,0.0002798497,0.003715122,0.00007741313,0.00028785726,0.0002295356,0.0005415139,0.7944274,0.004506754,0.030552482,0.16491993],"study_design_scores_gemma":[0.000047279103,0.00026203258,0.0010735234,0.00020672318,0.000066849694,0.00085153326,0.00006641967,0.0005314154,0.32872286,0.0025282826,0.66557026,0.000072881194],"about_ca_topic_score_codex":0.00043256945,"about_ca_topic_score_gemma":0.0010212273,"teacher_disagreement_score":0.031576775,"about_ca_system_score_codex":0.0005976534,"about_ca_system_score_gemma":0.0006253002,"threshold_uncertainty_score":0.10563487},"labels":[],"label_agreement":null},{"id":"W4232863667","doi":"10.3115/1567619.1567636","title":"Postnominal prepositional phrase attachment in proteomics","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Heuristics; Computer science; Noun phrase; Natural language processing; Phrase; Nominalization; Set (abstract data type); Artificial intelligence; Test set; Information retrieval; Programming language; Noun","score_opus":0.005977868209777394,"score_gpt":0.25102877536951435,"score_spread":0.24505090715973696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232863667","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17864259,0.0037473592,0.77950966,0.0012222342,0.00061643845,0.00069661636,0.0052445615,0.015398578,0.014922038],"genre_scores_gemma":[0.5048911,0.0014088298,0.4779651,0.00048579974,0.00025634613,0.00024851534,0.008297448,0.0009985708,0.005448228],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9978186,0.00077256653,0.0003076999,0.0005491915,0.00042775422,0.00012413763],"domain_scores_gemma":[0.99026066,0.0064630765,0.0010119815,0.0009266537,0.0011109145,0.0002266426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029887268,0.0011501438,0.00078949874,0.0030561266,0.0014327958,0.0024791723,0.0011841852,0.0011927567,0.005757216],"category_scores_gemma":[0.0095309755,0.00076251256,0.00082378276,0.0027785536,0.0010953717,0.0035045291,0.0014402145,0.0013688633,0.004552203],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016913869,0.00034201218,0.023588467,0.003914632,0.00018703028,0.002669711,0.0059919455,0.0071002105,0.16498011,0.04111386,0.035060946,0.7133598],"study_design_scores_gemma":[0.00026990252,0.00079451804,0.065584704,0.0012364719,0.00056092313,0.007927815,0.007061941,0.22226653,0.2968627,0.1339928,0.26295015,0.00049157214],"about_ca_topic_score_codex":0.0017318513,"about_ca_topic_score_gemma":0.0027782538,"teacher_disagreement_score":0.005757216,"about_ca_system_score_codex":0.0008364246,"about_ca_system_score_gemma":0.0014280853,"threshold_uncertainty_score":0.01925981},"labels":[],"label_agreement":null},{"id":"W4232955428","doi":"10.22215/etd/2014-10530","title":"Semantic Approaches to Enable Drug Discovery in Biomedical Big Data","year":2014,"lang":"en","type":"dissertation","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Drug repositioning; Computer science; Profiling (computer programming); Data science; Drug; Drug discovery; Data quality; Repurposing; Medicine; Bioinformatics; Pharmacology; Engineering; Biology","score_opus":0.10889259956627212,"score_gpt":0.292124389351426,"score_spread":0.18323178978515386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4232955428","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006042698,0.017049586,0.936491,0.019343361,0.00075274816,0.00037679676,0.0024064737,0.0015263013,0.016011098],"genre_scores_gemma":[0.079011306,0.020649888,0.8875009,0.0026211864,0.00065570924,0.00047368658,0.0052849883,0.00020068059,0.0036017436],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960503,0.0017730158,0.0004620861,0.00042515952,0.0011568159,0.00013260084],"domain_scores_gemma":[0.99313205,0.0043967213,0.00032865859,0.0012857396,0.0006480358,0.00020875713],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008334844,0.0009518783,0.0011526176,0.0043180855,0.0012994943,0.005773222,0.0015257056,0.0014375473,0.0030656252],"category_scores_gemma":[0.012232964,0.00060500397,0.002764343,0.0064625465,0.0023732607,0.010087915,0.0054250057,0.0036082682,0.0010431366],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008340995,0.00018421542,0.001251595,0.001486126,0.00042700738,0.00023282038,0.00073769386,0.012981345,0.0023708735,0.7436025,0.023157755,0.21348472],"study_design_scores_gemma":[0.000021990345,0.000022975506,0.00055661163,0.00031910307,0.00009764232,0.00011032755,0.00037734112,0.04377124,0.0018079412,0.85659105,0.09629464,0.00002921344],"about_ca_topic_score_codex":0.0020180861,"about_ca_topic_score_gemma":0.003667247,"teacher_disagreement_score":0.008334844,"about_ca_system_score_codex":0.0019848524,"about_ca_system_score_gemma":0.0030682716,"threshold_uncertainty_score":0.044079423},"labels":[],"label_agreement":null},{"id":"W4233240602","doi":"10.4018/978-1-5225-5191-1.ch032","title":"A Decision Support System (DSS) for Colorectal Cancer Follow-Up Program via a Semantic Framework","year":2018,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Laurentian University","funders":"","keywords":"Ontology; Decision support system; Flexibility (engineering); Web application; Semantic Web; Colorectal cancer; Interface (matter); Medicine; Computer science; User interface; Clinical decision support system; World Wide Web; Cancer; Artificial intelligence; Internal medicine","score_opus":0.01683010046051988,"score_gpt":0.3015032268766403,"score_spread":0.2846731264161204,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233240602","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012480465,0.00044166922,0.96963227,0.0013814669,0.00012901177,0.00038072158,0.0007414004,0.006224817,0.008588187],"genre_scores_gemma":[0.176454,0.00077760033,0.81639665,0.0004066578,0.000058459736,0.00029979894,0.0017480489,0.00026687828,0.003591922],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985405,0.000411065,0.00025523824,0.00024630062,0.0004265689,0.000120300014],"domain_scores_gemma":[0.99900764,0.0003971097,0.00008474231,0.00013383708,0.00027848058,0.00009825631],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025066193,0.000553239,0.0005091286,0.00185561,0.0008931422,0.0038929603,0.0014186146,0.0012780877,0.0017968041],"category_scores_gemma":[0.0027990604,0.00033489196,0.0012031198,0.0013073768,0.0008593152,0.0030119482,0.0014464332,0.001000829,0.0007624712],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062573346,0.0006344358,0.0060490854,0.0013270911,0.00028659738,0.0019511225,0.00278294,0.10077249,0.026712459,0.43441665,0.02483524,0.3996062],"study_design_scores_gemma":[0.00016736076,0.00024507573,0.002019026,0.0006238706,0.0003226825,0.0009663116,0.0010779239,0.5099589,0.022692235,0.1289394,0.3328028,0.00018444534],"about_ca_topic_score_codex":0.0074627902,"about_ca_topic_score_gemma":0.00642944,"teacher_disagreement_score":0.0074627902,"about_ca_system_score_codex":0.0011893961,"about_ca_system_score_gemma":0.0028278155,"threshold_uncertainty_score":0.0148386955},"labels":[],"label_agreement":null},{"id":"W4233525735","doi":"10.1002/asi.21414","title":"Adapting semantic natural language processing technology to address information overload in influenza epidemic management","year":2010,"lang":"en","type":"article","venue":"Journal of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"National Institutes of Health; New York Academy of Medicine","keywords":"Automatic summarization; Computer science; Information overload; Terminology; Ontology; Semantics (computer science); Information retrieval; Semantic analysis (machine learning); Natural language processing; Data science; Artificial intelligence; World Wide Web; Linguistics","score_opus":0.007414650250070422,"score_gpt":0.311415120456833,"score_spread":0.3040004702067626,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233525735","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10697279,0.00017071972,0.8766689,0.0017701539,0.00008470467,0.0006885525,0.00035577823,0.009203019,0.0040854937],"genre_scores_gemma":[0.24973355,0.00020702725,0.74742717,0.00042979285,0.000060040496,0.00035981362,0.0005881158,0.0002755462,0.0009189933],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9962155,0.0024350851,0.0003522045,0.0003512478,0.00056815473,0.0000777735],"domain_scores_gemma":[0.98291147,0.0135719525,0.00053193275,0.0012809004,0.0015736752,0.00013006566],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0077977944,0.0005496764,0.00053292373,0.001913834,0.0005429853,0.0021551412,0.0010544785,0.0008586461,0.0019112198],"category_scores_gemma":[0.021530919,0.00030933355,0.00062308746,0.0014149158,0.00075284473,0.004645328,0.0018248464,0.00094888476,0.00066259765],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012298595,0.0013446334,0.008003963,0.0021284365,0.00030109024,0.0019026877,0.010690658,0.03714504,0.08739183,0.029246602,0.0145448055,0.8060703],"study_design_scores_gemma":[0.00042701646,0.000674369,0.005018482,0.00040194707,0.000496911,0.0012572892,0.0042148447,0.7120622,0.09107683,0.10516096,0.079014465,0.00019474483],"about_ca_topic_score_codex":0.0007780022,"about_ca_topic_score_gemma":0.00094489224,"teacher_disagreement_score":0.0077977944,"about_ca_system_score_codex":0.0005668301,"about_ca_system_score_gemma":0.0008413414,"threshold_uncertainty_score":0.041239142},"labels":[],"label_agreement":null},{"id":"W4233767670","doi":"10.1007/978-3-658-21155-4_6","title":"String-Algorithmen","year":2018,"lang":"de","type":"book-chapter","venue":"Studienbücher Informatik","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.02499907116265648,"score_gpt":0.2686512060208274,"score_spread":0.24365213485817094,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4233767670","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004229582,0.0011550314,0.927376,0.0011483656,0.000928215,0.00041950756,0.005767201,0.024667414,0.03430875],"genre_scores_gemma":[0.04448586,0.0015427535,0.87896967,0.0011045316,0.00038329326,0.001509736,0.023901803,0.00853465,0.03956777],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9942521,0.0012486845,0.00062478386,0.0016156511,0.0018923421,0.00036654127],"domain_scores_gemma":[0.9956398,0.0020382532,0.00013937376,0.0013278272,0.0007379445,0.000116919626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00263441,0.0036059064,0.0022562752,0.0046980646,0.0019764004,0.0074886717,0.0040721484,0.004486427,0.086233936],"category_scores_gemma":[0.015792683,0.002263747,0.0039080908,0.007149721,0.0021190695,0.009286307,0.006964587,0.009617174,0.08640129],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00048599578,0.00037235927,0.0005546108,0.0011688329,0.00027824982,0.00016861309,0.00021731252,0.038211286,0.0023475348,0.17485745,0.07895157,0.7023862],"study_design_scores_gemma":[0.0002338318,0.0001923637,0.00031351627,0.00036120982,0.00014197435,0.0004251157,0.00012479286,0.21559818,0.00826648,0.6631209,0.11113451,0.00008710558],"about_ca_topic_score_codex":0.0012644894,"about_ca_topic_score_gemma":0.0013779136,"teacher_disagreement_score":0.086233936,"about_ca_system_score_codex":0.0019092979,"about_ca_system_score_gemma":0.0025446392,"threshold_uncertainty_score":0.2884813},"labels":[],"label_agreement":null},{"id":"W4234416667","doi":"10.36227/techrxiv.12055998","title":"Drug-Drug Interaction Detection (DDI) Over the Social Media using Convolutional Neural Networks","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Convolutional neural network; Computer science; Drug-drug interaction; Drug; Artificial intelligence; Artificial neural network; Social media; Machine learning; Pharmacology; Medicine; World Wide Web","score_opus":0.03819318257508286,"score_gpt":0.30870577643930425,"score_spread":0.2705125938642214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234416667","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4632483,0.0124874385,0.4028049,0.005314663,0.000549591,0.00088914647,0.07432382,0.023225812,0.01715636],"genre_scores_gemma":[0.7357453,0.004080879,0.18897742,0.00047528086,0.00022474416,0.00032047028,0.058068413,0.00037319277,0.011734295],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993431,0.00012499162,0.000058054593,0.00018032778,0.00022006531,0.0000734715],"domain_scores_gemma":[0.9985934,0.0007653494,0.00018528514,0.00019059257,0.00021219032,0.000053324165],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009816006,0.00095882325,0.00048434,0.0032829237,0.00038957025,0.00085339166,0.00045026708,0.00067822216,0.0014795158],"category_scores_gemma":[0.002945391,0.00028373697,0.00082724745,0.0021438643,0.00024446368,0.0015431613,0.00085657253,0.00066268607,0.00090470666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00094850274,0.000765246,0.032606855,0.0013297753,0.001015312,0.00091432524,0.00029176945,0.0487745,0.04518706,0.007828127,0.04284672,0.8174918],"study_design_scores_gemma":[0.00007944658,0.0001822959,0.021549094,0.00016961133,0.0003787478,0.0004093479,0.0001349895,0.8916871,0.039327756,0.017378913,0.028650686,0.00005199715],"about_ca_topic_score_codex":0.0117440475,"about_ca_topic_score_gemma":0.020585146,"teacher_disagreement_score":0.0117440475,"about_ca_system_score_codex":0.0009227976,"about_ca_system_score_gemma":0.0012614201,"threshold_uncertainty_score":0.023351371},"labels":[],"label_agreement":null},{"id":"W4234604627","doi":"10.1093/ajcp/aqaa161.351","title":"What’s in a Name? Comparative Analysis of Laboratory Test Naming Guidelines as Applied to Common Confusing Test Names","year":2020,"lang":"en","type":"article","venue":"American Journal of Clinical Pathology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Respondent; Computer science","score_opus":0.08670043127588939,"score_gpt":0.44350665475596635,"score_spread":0.35680622348007696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4234604627","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9798545,0.0012408491,0.008291742,0.0026342724,0.00017952255,0.00081794034,0.00046591883,0.00011922,0.006395998],"genre_scores_gemma":[0.9867762,0.00091073493,0.009549369,0.0010203503,0.00004227496,0.00080528134,0.0004758243,0.000066574736,0.00035343098],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.8722026,0.07233637,0.021286799,0.0046864045,0.026037678,0.003450218],"domain_scores_gemma":[0.43492928,0.35083073,0.10487564,0.0123317,0.09229507,0.00473764],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09729452,0.00038390266,0.000768342,0.0071690376,0.0021167828,0.0053929887,0.0020177807,0.00159447,0.0016630098],"category_scores_gemma":[0.45792714,0.0004435386,0.0010404201,0.006533877,0.0039060847,0.007847427,0.0050181136,0.0017613684,0.00046138544],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008046384,0.00039754613,0.62001395,0.0032466364,0.00022483387,0.00073888694,0.2070672,0.0009611598,0.0022675628,0.0071165725,0.009885477,0.14727545],"study_design_scores_gemma":[0.00009714087,0.0010143165,0.5463592,0.004951677,0.00033298714,0.0010740939,0.39727554,0.008079409,0.0036701758,0.004018912,0.03278067,0.00034582117],"about_ca_topic_score_codex":0.013400098,"about_ca_topic_score_gemma":0.016572656,"teacher_disagreement_score":0.9027055,"about_ca_system_score_codex":0.0074900356,"about_ca_system_score_gemma":0.01294685,"threshold_uncertainty_score":0.5145488},"labels":[],"label_agreement":null},{"id":"W4235129037","doi":"10.1038/npre.2008.1784.1","title":"Suggested actions from the Melbourne HVP Information Seminar","year":2008,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Hôpital Notre-Dame","funders":"","keywords":"Library science; Political science; Computer science","score_opus":0.015536491484254381,"score_gpt":0.28408598256252354,"score_spread":0.26854949107826914,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235129037","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006566147,0.008958057,0.0019888217,0.5361407,0.31182718,0.00064211455,0.0018786833,0.001550455,0.13635738],"genre_scores_gemma":[0.00719562,0.003365579,0.0030729019,0.18995553,0.082656875,0.0006252915,0.0017501756,0.0010402896,0.71033776],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99186265,0.0018649452,0.00052124343,0.0009493703,0.0033181089,0.0014836811],"domain_scores_gemma":[0.9813348,0.0031331377,0.00095499365,0.00068041997,0.0050842073,0.008812488],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.013010452,0.0021631534,0.0012003308,0.0031446067,0.0041011833,0.01126423,0.00415458,0.023947323,0.30391252],"category_scores_gemma":[0.027786102,0.001098382,0.0019293224,0.0016622309,0.0014919286,0.007260557,0.009236197,0.017772056,0.15299821],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022010392,0.000014929153,0.000015352782,0.00004645331,0.0000013633555,0.00008540796,0.000037405178,0.000014910684,0.00008142676,0.0010309096,0.9938235,0.0048261257],"study_design_scores_gemma":[0.000015243522,0.000015338892,0.00017273429,0.00007796264,0.000003168594,0.00005309715,0.00009414876,0.00003510016,0.00006843854,0.0006489305,0.99879944,0.00001637512],"about_ca_topic_score_codex":0.0065094084,"about_ca_topic_score_gemma":0.009542935,"teacher_disagreement_score":0.30391252,"about_ca_system_score_codex":0.0054225978,"about_ca_system_score_gemma":0.008117764,"threshold_uncertainty_score":0.99288434},"labels":[],"label_agreement":null},{"id":"W4235165249","doi":"10.5663/aps.v7i1.29345","title":"Network Environments for Aboriginal Health Research (NEAHR)","year":2018,"lang":"en","type":"article","venue":"aboriginal policy studies","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Geography","score_opus":0.13664108797422597,"score_gpt":0.5135568636776,"score_spread":0.37691577570337403,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235165249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029328471,0.0035463504,0.73941284,0.016132588,0.001428069,0.0013363616,0.043022335,0.054233853,0.11155912],"genre_scores_gemma":[0.09521644,0.0031559074,0.81692535,0.0015612091,0.00039176844,0.0012695785,0.043210078,0.0038801096,0.03438959],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978416,0.001184292,0.00019980394,0.00027120623,0.00037809316,0.00012484097],"domain_scores_gemma":[0.98762053,0.0069369883,0.0005423296,0.0026153445,0.0011108547,0.0011739893],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009768527,0.0004844107,0.0003322161,0.0017525319,0.00144604,0.002737412,0.001329496,0.0007721349,0.011586468],"category_scores_gemma":[0.015382586,0.00052532065,0.0008335007,0.0018376171,0.0006922784,0.0036845992,0.006434315,0.0011743743,0.0032807793],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00054534647,0.00032312926,0.011677882,0.0027540768,0.00028126806,0.00086993334,0.011319453,0.009858045,0.0075208596,0.3410262,0.24044028,0.3733836],"study_design_scores_gemma":[0.000047293277,0.000053004464,0.0048628454,0.00046986208,0.000090436704,0.00028933454,0.0012603701,0.012069843,0.0027001423,0.11408822,0.8640131,0.000055519864],"about_ca_topic_score_codex":0.012539475,"about_ca_topic_score_gemma":0.023208322,"teacher_disagreement_score":0.98746055,"about_ca_system_score_codex":0.0010439929,"about_ca_system_score_gemma":0.003637718,"threshold_uncertainty_score":0.05166155},"labels":[],"label_agreement":null},{"id":"W4235282642","doi":"10.21203/rs.3.rs-40780/v2","title":"Decoding semi-automated title-abstract screening: a retrospective exploration of the review, study, and publication characteristics associated with accurate relevance predictions","year":2020,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Relevance (law); Decoding methods; Information retrieval; Computer science; Data science; Algorithm; Political science","score_opus":0.10737056613387296,"score_gpt":0.39861958659113605,"score_spread":0.2912490204572631,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4235282642","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.63176167,0.052054413,0.25238547,0.005601855,0.0006803422,0.006449043,0.044212695,0.0032272027,0.0036273152],"genre_scores_gemma":[0.8945676,0.003115374,0.08756138,0.001032999,0.00018337835,0.003659631,0.009143826,0.00037967172,0.0003562243],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.73383915,0.18600154,0.049849942,0.012381509,0.016690977,0.0012368852],"domain_scores_gemma":[0.15398663,0.7154519,0.08192882,0.030178972,0.017642748,0.00081081165],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2326506,0.0015245185,0.002448966,0.006215493,0.0006355573,0.003254774,0.0017886389,0.001167142,0.0025596167],"category_scores_gemma":[0.639156,0.0014242594,0.006765517,0.0072985897,0.001081927,0.0034889802,0.0023653293,0.0014939962,0.0011871674],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009649062,0.00024757654,0.5755705,0.03581409,0.0152816875,0.0013478096,0.002737292,0.022906218,0.004313636,0.0019164953,0.016655315,0.3135603],"study_design_scores_gemma":[0.0052406425,0.0075472663,0.5412112,0.025617586,0.06440071,0.0109184645,0.0019135317,0.22919215,0.025456,0.023461165,0.063755564,0.0012857607],"about_ca_topic_score_codex":0.0021133109,"about_ca_topic_score_gemma":0.0034882163,"teacher_disagreement_score":0.7673494,"about_ca_system_score_codex":0.0017781791,"about_ca_system_score_gemma":0.005148071,"threshold_uncertainty_score":0.94627845},"labels":[],"label_agreement":null},{"id":"W4236062438","doi":"10.1038/npre.2009.3635.1","title":"Hematopoietic Cell Types: Prototype for a Revised Cell Ontology","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institute of Allergy and Infectious Diseases; National Human Genome Research Institute","keywords":"Ontology; Computer science; Process ontology; Open Biomedical Ontologies; Ontology-based data integration; Suggested Upper Merged Ontology; Gene ontology; Upper ontology; Structuring; Hematopoietic cell; Representation (politics); Computational biology; Information retrieval; Haematopoiesis; Chemistry; Biology; Gene; Stem cell; Biochemistry; Cell biology; Semantic Web","score_opus":0.011009392607371339,"score_gpt":0.29068169881707195,"score_spread":0.2796723062097006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4236062438","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016630042,0.0004146305,0.92385304,0.0048299036,0.0007616773,0.001564177,0.0066805934,0.028615378,0.016650505],"genre_scores_gemma":[0.055030283,0.0004910287,0.9107503,0.0010279686,0.00011325284,0.00051193824,0.017979767,0.003020665,0.011074764],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99610496,0.0006472978,0.0006634036,0.0006568695,0.0016662473,0.00026122882],"domain_scores_gemma":[0.993222,0.002176504,0.0003224807,0.0017163258,0.0020369291,0.00052567513],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008532015,0.00043716916,0.0007038973,0.0025785717,0.0014397838,0.0053917333,0.004203463,0.00196793,0.006893241],"category_scores_gemma":[0.01255464,0.0008244963,0.0020557882,0.0024852538,0.0014804249,0.0075995373,0.004380972,0.002948765,0.0032934302],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082429877,0.0010378317,0.006256097,0.0018761572,0.00018785203,0.0029370924,0.005507844,0.011784133,0.042163376,0.37528077,0.17526792,0.37687653],"study_design_scores_gemma":[0.00019880907,0.00010897632,0.0013618863,0.0004287943,0.00012494887,0.0012105101,0.0008766256,0.06594487,0.016611421,0.048164524,0.86486363,0.00010499276],"about_ca_topic_score_codex":0.011900817,"about_ca_topic_score_gemma":0.014525896,"teacher_disagreement_score":0.011900817,"about_ca_system_score_codex":0.0023753333,"about_ca_system_score_gemma":0.0036917024,"threshold_uncertainty_score":0.045122147},"labels":[],"label_agreement":null},{"id":"W4239223488","doi":"10.1503/cmaj.1040382","title":"Corrections","year":2004,"lang":"en","type":"article","venue":"Canadian Medical Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Data science; Information retrieval; World Wide Web","score_opus":0.006272678842282943,"score_gpt":0.23677071303614045,"score_spread":0.2304980341938575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239223488","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00025376503,0.0010496883,0.0032392305,0.124189384,0.8487914,0.00013774166,0.0036755328,0.0015647389,0.017098498],"genre_scores_gemma":[0.014655978,0.0048656464,0.014806863,0.15870163,0.23982108,0.0006922588,0.008318105,0.0057447497,0.55239373],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9838489,0.0029841461,0.0023492272,0.0017374477,0.008135577,0.0009446281],"domain_scores_gemma":[0.87257457,0.025708036,0.005122312,0.00689994,0.08671084,0.0029843296],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008401887,0.0019188429,0.0011293757,0.0053823134,0.003103644,0.0052158316,0.0032668733,0.0048188134,0.18067054],"category_scores_gemma":[0.17265809,0.0010392106,0.0018496554,0.0031621726,0.0020833763,0.00322819,0.0027457334,0.008871112,0.11127718],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015433343,0.0000023919968,0.000039594906,0.000047094472,0.000003941792,0.000062828716,0.00003152635,0.000018607821,0.000030283947,0.0010025171,0.99269426,0.0060514463],"study_design_scores_gemma":[0.000016262076,0.0000081359985,0.00028750574,0.00016878636,0.000010432662,0.00023411821,0.00006589949,0.00007003541,0.00010895432,0.0013858603,0.99762684,0.000017137594],"about_ca_topic_score_codex":0.014390794,"about_ca_topic_score_gemma":0.020140205,"teacher_disagreement_score":0.18067054,"about_ca_system_score_codex":0.0052629192,"about_ca_system_score_gemma":0.007591587,"threshold_uncertainty_score":0.6044032},"labels":[],"label_agreement":null},{"id":"W4239370165","doi":"10.1201/b13123-5","title":"The Availability of MeSH in Vendor-Supplied Cataloguing Records, as Seen Through the Catalogue of a Canadian Academic Health Library","year":2011,"lang":"en","type":"book-chapter","venue":"Apple Academic Press eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Vendor; Library science; Computer science; World Wide Web; Information retrieval; Database; Business; Marketing","score_opus":0.05239443820699945,"score_gpt":0.27597492278744684,"score_spread":0.22358048458044738,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239370165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.61247534,0.023567058,0.00803534,0.007788041,0.0005492164,0.00066716055,0.12153794,0.0019959891,0.223384],"genre_scores_gemma":[0.83090943,0.024610167,0.017485965,0.0017984691,0.00026268046,0.00031034794,0.085914075,0.0009212918,0.03778762],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9725381,0.0019220173,0.0037574773,0.0013770068,0.019159744,0.0012455967],"domain_scores_gemma":[0.8030719,0.060945503,0.045166243,0.01270217,0.07337325,0.004740908],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.011610087,0.00035241342,0.00060576154,0.05611694,0.0036547172,0.008026108,0.0020394167,0.0005980133,0.009775869],"category_scores_gemma":[0.11602284,0.0005784142,0.0004350504,0.116180405,0.001966612,0.004508945,0.0035893589,0.00063247373,0.0033025527],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027259524,0.00004961907,0.5697521,0.004124124,0.00019156483,0.0013292365,0.032465413,0.0002442269,0.004347091,0.013801899,0.082041845,0.29138035],"study_design_scores_gemma":[0.000018335391,0.00006086147,0.61244833,0.0020481145,0.00022115746,0.0018538915,0.016800806,0.00034754316,0.0023972476,0.0009015091,0.36273092,0.00017129618],"about_ca_topic_score_codex":0.58835155,"about_ca_topic_score_gemma":0.6437506,"teacher_disagreement_score":0.9919739,"about_ca_system_score_codex":0.011791526,"about_ca_system_score_gemma":0.034247037,"threshold_uncertainty_score":0.82814544},"labels":[],"label_agreement":null},{"id":"W4239774680","doi":"10.1111/j.1469-7580.2005.00490.x","title":"BOOK REVIEW","year":2005,"lang":"en","type":"article","venue":"Journal of Anatomy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Anatomy; Medicine","score_opus":0.007630588239579297,"score_gpt":0.3012072517612679,"score_spread":0.2935766635216886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4239774680","genre_codex":"other","genre_gemma":"editorial","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0007276747,0.17039561,0.0026071786,0.01972496,0.03906812,0.00024428457,0.004727997,0.0009935412,0.7615106],"genre_scores_gemma":[0.0033308119,0.08050342,0.0020287393,0.008090173,0.0088379355,0.0001264662,0.003924318,0.0005191381,0.8926391],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99912554,0.00009717622,0.0000678609,0.00017050974,0.0004584643,0.000080460944],"domain_scores_gemma":[0.99826044,0.00032995458,0.000092842616,0.00015778307,0.00085906876,0.00029987807],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00060104916,0.0008801856,0.0012167987,0.003316039,0.0012777665,0.0059700664,0.002213142,0.0021562479,0.4883438],"category_scores_gemma":[0.004376384,0.00038888492,0.0007177028,0.0040045357,0.0007160454,0.0040408745,0.002229379,0.0024470566,0.39278615],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00001579776,0.000017972057,0.00006504391,0.00044519786,0.0000054750144,0.00010098054,0.00006975938,0.000041230516,0.000102320635,0.0041479105,0.85835475,0.13663357],"study_design_scores_gemma":[0.0000019426766,0.0000035322287,0.000066614906,0.00020062027,0.0000021004764,0.00017784859,0.000031248823,0.0000079061565,0.000018110839,0.0006049842,0.9988827,0.000002450975],"about_ca_topic_score_codex":0.0026567553,"about_ca_topic_score_gemma":0.0046836385,"teacher_disagreement_score":0.4883438,"about_ca_system_score_codex":0.0017128388,"about_ca_system_score_gemma":0.0033681714,"threshold_uncertainty_score":0.7298155},"labels":[],"label_agreement":null},{"id":"W4240553231","doi":"10.5596/c2012-018","title":"Canadian Virtual Health Library / Bibliothèque virtuelle canadienne de la santé (CVHL/BVCS)","year":2012,"lang":"fr","type":"article","venue":"Journal of the Canadian Health Libraries Association / Journal de l Association de bilbiothèques de la santé du Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Library science; Art; Computer science","score_opus":0.004396205842778417,"score_gpt":0.24544877823706993,"score_spread":0.2410525723942915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4240553231","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036508585,0.015143559,0.0063544414,0.08271691,0.0061078044,0.00052229565,0.10940805,0.0042698737,0.7718262],"genre_scores_gemma":[0.095354676,0.02777255,0.038183104,0.01726954,0.0021671278,0.0006479558,0.10705225,0.0023395284,0.7092132],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99047977,0.0010412711,0.0005348378,0.0005586001,0.0062923445,0.001093173],"domain_scores_gemma":[0.9648607,0.0036144939,0.0012612888,0.0019357908,0.019847343,0.008480406],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.007092059,0.0008755169,0.0010748556,0.011228662,0.00798117,0.015628861,0.0027105336,0.0028308695,0.14138377],"category_scores_gemma":[0.032284018,0.00079751725,0.0008091481,0.017018335,0.0024039587,0.003538335,0.004303266,0.0028624036,0.0211297],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030617615,0.000021100654,0.0009191187,0.00025158827,0.00001729145,0.000041819036,0.000115980685,0.00016325619,0.00010415008,0.0256056,0.9165708,0.05615873],"study_design_scores_gemma":[0.000013785083,0.000003657159,0.004004094,0.00015785042,0.000011551937,0.000037110207,0.00011814729,0.00024676035,0.00013911995,0.002362455,0.9928803,0.000025267966],"about_ca_topic_score_codex":0.9444306,"about_ca_topic_score_gemma":0.9387705,"teacher_disagreement_score":0.9843711,"about_ca_system_score_codex":0.06313882,"about_ca_system_score_gemma":0.25177896,"threshold_uncertainty_score":0.4729758},"labels":[],"label_agreement":null},{"id":"W4241866776","doi":"10.31219/osf.io/zep3x","title":"Capturing scientific knowledge in computable form","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Open Knowledge Base Connectivity; Knowledge representation and reasoning; Context (archaeology); Knowledge management; Data science; Sociology of scientific knowledge; Process (computing); Knowledge extraction; Procedural knowledge; Domain knowledge; Personal knowledge management; Artificial intelligence; Organizational learning","score_opus":0.046026872251764196,"score_gpt":0.30610895023865875,"score_spread":0.26008207798689453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4241866776","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010279262,0.0008712041,0.9656673,0.0020312725,0.00020287701,0.00014924571,0.0026913271,0.0022294743,0.015878068],"genre_scores_gemma":[0.14428844,0.0026380508,0.83495355,0.0004921574,0.00029832413,0.0004750396,0.009931888,0.0006060013,0.0063164593],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9943499,0.0013900277,0.00079168926,0.0008913884,0.0023254997,0.00025146818],"domain_scores_gemma":[0.9866838,0.0070984876,0.0008273227,0.0035896858,0.0015603653,0.00024027652],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039071906,0.001149979,0.0011792526,0.006981357,0.0014999241,0.013822113,0.00218763,0.0018963838,0.007210211],"category_scores_gemma":[0.03511994,0.0008802287,0.0021916516,0.009455519,0.0039670398,0.015405909,0.0067898137,0.0024875896,0.0029084242],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009588416,0.00006989026,0.0016249407,0.0008612704,0.00010872157,0.0006774504,0.0016328165,0.039945472,0.002743508,0.8178953,0.009932315,0.12441242],"study_design_scores_gemma":[0.000026484511,0.00001786457,0.00028035406,0.00019843977,0.0000634271,0.00017909204,0.0003565356,0.053875424,0.0029148334,0.86601615,0.07604176,0.000029625373],"about_ca_topic_score_codex":0.0046468186,"about_ca_topic_score_gemma":0.004433837,"teacher_disagreement_score":0.013822113,"about_ca_system_score_codex":0.002118023,"about_ca_system_score_gemma":0.0040157307,"threshold_uncertainty_score":0.024120629},"labels":[],"label_agreement":null},{"id":"W4242036144","doi":"10.5858/134.5.663","title":"Electronic Pathology Reporting: Digitizing the College of American Pathologists Cancer Checklists","year":2010,"lang":"en","type":"letter","venue":"Archives of Pathology & Laboratory Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Checklist; Accreditation; Medicine; Cancer; Commission; Family medicine; Health care; Medical physics; MEDLINE; Medical education; Pathology; Psychology; Internal medicine; Business; Political science","score_opus":0.015341084953277939,"score_gpt":0.30104012260064905,"score_spread":0.2856990376473711,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242036144","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015783093,0.064038545,0.06014418,0.35529575,0.49313313,0.0008913299,0.00291285,0.0036733767,0.018332535],"genre_scores_gemma":[0.01434959,0.1578362,0.19929054,0.14616449,0.44917676,0.0023940993,0.009228799,0.003189958,0.018369444],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.93963045,0.022412281,0.01277475,0.0022830185,0.021657847,0.0012416617],"domain_scores_gemma":[0.7197114,0.11960804,0.021392792,0.015517244,0.11548011,0.008290378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.052499305,0.0012705529,0.0013291523,0.017629974,0.0019893413,0.010789229,0.006088472,0.004355321,0.010707458],"category_scores_gemma":[0.22111659,0.0013636495,0.0013548816,0.011070244,0.003730763,0.013580675,0.006154357,0.009894279,0.0077926787],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000023463816,0.000018854784,0.00051455526,0.0007638963,0.000021391366,0.000046556597,0.0002592722,0.00013381237,0.0001699052,0.004369127,0.80201924,0.19165994],"study_design_scores_gemma":[0.000012049852,0.000018584351,0.00043277448,0.0012799281,0.000018874493,0.00012333342,0.00013922335,0.00020883445,0.00021157238,0.0018980904,0.99562067,0.000036130376],"about_ca_topic_score_codex":0.004638558,"about_ca_topic_score_gemma":0.0072329966,"teacher_disagreement_score":0.052499305,"about_ca_system_score_codex":0.0042989594,"about_ca_system_score_gemma":0.009817444,"threshold_uncertainty_score":0.27764618},"labels":[],"label_agreement":null},{"id":"W4242777923","doi":"10.17504/protocols.io.bui7nuhn","title":"SARS-CoV-2 NCBI submission protocol: SRA, BioSample, and BioProject v3","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Protocol (science); Biology; Coronavirus disease 2019 (COVID-19); Metadata; Computational biology; World Wide Web; Computer science; Medicine; Pathology","score_opus":0.054263752994182415,"score_gpt":0.3599418889722303,"score_spread":0.3056781359780479,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242777923","genre_codex":"dataset","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077780313,0.0037148586,0.17157882,0.009019031,0.009547637,0.03744858,0.5893143,0.06561219,0.10598648],"genre_scores_gemma":[0.010880839,0.002536483,0.10904632,0.006438691,0.0014008705,0.033542335,0.72904223,0.024154799,0.08295741],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97950965,0.0071545336,0.0038102672,0.0024215074,0.005396293,0.0017078753],"domain_scores_gemma":[0.970834,0.0057213334,0.0020336108,0.008223017,0.01131898,0.0018689316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019103443,0.0023284392,0.0028192808,0.0051439824,0.005523264,0.00564935,0.0050945855,0.004628121,0.25062984],"category_scores_gemma":[0.038311314,0.0033544607,0.0021096647,0.0043292427,0.0017409364,0.005197671,0.0063095703,0.0053596417,0.34811124],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009582244,0.00016830982,0.0010349109,0.0019104178,0.000049475166,0.0001878461,0.0003605725,0.00013516402,0.02432088,0.0032321606,0.94702655,0.02061552],"study_design_scores_gemma":[0.00029601637,0.00031691723,0.0023397515,0.0008906542,0.000045763973,0.00037815244,0.00032083096,0.00032265912,0.022808425,0.0026957546,0.9694437,0.00014132532],"about_ca_topic_score_codex":0.0044743763,"about_ca_topic_score_gemma":0.006562518,"teacher_disagreement_score":0.25062984,"about_ca_system_score_codex":0.0021956374,"about_ca_system_score_gemma":0.009402911,"threshold_uncertainty_score":0.83844036},"labels":[],"label_agreement":null},{"id":"W4242805031","doi":"10.4018/978-1-4666-1755-1.ch006","title":"A Framework for Data and Mined Knowledge Interoperability in Clinical Decision Support Systems","year":2012,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Interoperability; Clinical decision support system; Decision support system; Guideline; Computer science; Knowledge management; Health care; Quality (philosophy); Data science; Data mining; Medicine; World Wide Web","score_opus":0.08511778478953548,"score_gpt":0.38800112265137326,"score_spread":0.3028833378618378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242805031","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00042618616,0.0010938166,0.9893811,0.0016283884,0.000092003815,0.00025423133,0.00032600953,0.0009355356,0.005862788],"genre_scores_gemma":[0.007819338,0.0009456317,0.98717505,0.00036629947,0.000073800016,0.0004750185,0.0009374905,0.00013991032,0.0020674015],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99030346,0.0040537375,0.0017829846,0.0010738702,0.0024170296,0.00036898698],"domain_scores_gemma":[0.9935714,0.0038449806,0.0003429068,0.001337104,0.0006010034,0.00030267626],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015790237,0.001649458,0.00164826,0.005476495,0.0023866391,0.0141568845,0.00771994,0.005459567,0.00661346],"category_scores_gemma":[0.015073315,0.0017846649,0.0043984796,0.00873544,0.0055408026,0.013430157,0.008945193,0.006465729,0.0032031378],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025847386,0.00006093111,0.00022748075,0.0005805532,0.00007752408,0.00058913603,0.0011783491,0.01017713,0.000938105,0.8847858,0.009082321,0.09227676],"study_design_scores_gemma":[0.000035110468,0.000036480676,0.0001458395,0.00071735034,0.00005718313,0.0005004098,0.0003531769,0.063336,0.0013652402,0.7371114,0.1962859,0.00005582126],"about_ca_topic_score_codex":0.008258087,"about_ca_topic_score_gemma":0.007619721,"teacher_disagreement_score":0.015790237,"about_ca_system_score_codex":0.004066608,"about_ca_system_score_gemma":0.0054980177,"threshold_uncertainty_score":0.083507776},"labels":[],"label_agreement":null},{"id":"W4242907897","doi":"10.1038/npre.2010.4270.1","title":"Formulating MEDLINE queries for article retrieval based on PubMed exemplars","year":2010,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Institutes of Health","keywords":"Computer science; Information retrieval; Search engine indexing; Set (abstract data type); Bigram; Task (project management); Result set; Process (computing); Function (biology); Recall; Natural language processing","score_opus":0.014892616225833525,"score_gpt":0.29648085099582633,"score_spread":0.2815882347699928,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4242907897","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12663923,0.003213281,0.7069574,0.0059507843,0.00039932784,0.003926329,0.047528982,0.09200779,0.013376874],"genre_scores_gemma":[0.12772241,0.0010887196,0.83283365,0.00058345357,0.00017841636,0.00096158136,0.03266838,0.0017806554,0.0021827044],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954383,0.0010536843,0.0012458146,0.00067678874,0.0013525705,0.00023290365],"domain_scores_gemma":[0.9804709,0.014344707,0.001169289,0.0009894549,0.0024822059,0.0005434055],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.003857875,0.0019336087,0.0023108523,0.011916991,0.0013339403,0.0044674594,0.0018617972,0.0025042754,0.013060763],"category_scores_gemma":[0.031715255,0.0010972196,0.0022263762,0.006463361,0.00075608765,0.006328223,0.002942056,0.0011501138,0.0053527374],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0034168933,0.0007931826,0.019090397,0.010376399,0.00065198774,0.004547473,0.004273445,0.018254835,0.09683853,0.045420606,0.13467988,0.6616564],"study_design_scores_gemma":[0.0015091241,0.0011999988,0.014796443,0.001432851,0.00094010757,0.005262696,0.0049531683,0.49648836,0.12191506,0.086901285,0.26408353,0.0005173482],"about_ca_topic_score_codex":0.0035192708,"about_ca_topic_score_gemma":0.0060635107,"teacher_disagreement_score":0.99614215,"about_ca_system_score_codex":0.0017969743,"about_ca_system_score_gemma":0.0018890917,"threshold_uncertainty_score":0.04369253},"labels":[],"label_agreement":null},{"id":"W4243589020","doi":"10.32388/isaaws","title":"CDISC SDTM Canadian Cardiovascular Society Classification Terminology","year":2020,"lang":"en","type":"reference-entry","venue":"Definitions","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Terminology; Computer science; Data science; Information retrieval; Medicine; Linguistics","score_opus":0.0997321864611979,"score_gpt":0.2705333523112378,"score_spread":0.1708011658500399,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243589020","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014809272,0.0019384504,0.03935691,0.0034815723,0.002740463,0.00057123345,0.8099702,0.010286646,0.13017367],"genre_scores_gemma":[0.0070236675,0.0026118706,0.05626948,0.0012234928,0.000402989,0.0006614644,0.87796104,0.0026112795,0.0512347],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99691737,0.00036060953,0.0007673659,0.00047620758,0.0011711784,0.00030730764],"domain_scores_gemma":[0.9877329,0.0017754459,0.000704935,0.0012919416,0.007933617,0.00056104537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023311414,0.0017367541,0.0011384116,0.015941272,0.0020230368,0.0050725685,0.0031684118,0.0016370051,0.103687055],"category_scores_gemma":[0.016175577,0.00047233186,0.0009272431,0.019750899,0.0009550487,0.005017147,0.0028894888,0.0020615906,0.08202534],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035393867,0.000009784976,0.00033691496,0.00047801365,0.000004644325,0.000037746777,0.00007189634,0.00026673477,0.00034589288,0.02160007,0.9335169,0.043295935],"study_design_scores_gemma":[0.000008026945,0.000003296033,0.00030021058,0.00022258113,0.000008576379,0.00004917517,0.00005248856,0.00050451263,0.0004129197,0.0040274905,0.9943962,0.000014487561],"about_ca_topic_score_codex":0.17654212,"about_ca_topic_score_gemma":0.1147284,"teacher_disagreement_score":0.8234579,"about_ca_system_score_codex":0.0057074195,"about_ca_system_score_gemma":0.024273576,"threshold_uncertainty_score":0.35102904},"labels":[],"label_agreement":null},{"id":"W4243745316","doi":"10.3410/f.740214753.793586320","title":"Faculty Opinions recommendation of A proximity-dependent biotinylation map of a human cell.","year":2021,"lang":"en","type":"dataset","venue":"Faculty Opinions – Post-Publication Peer Review of the Biomedical Literature","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval","funders":"","keywords":"Biotinylation; Proteome; Endoplasmic reticulum; Compartmentalization (fire protection); Organelle; Computational biology; Subcellular localization; Intracellular; Cell biology; Biology; Cellular compartment; Cell; Mitochondrion; Computer science; Bioinformatics; Biochemistry; Cytoplasm","score_opus":0.03182545596660652,"score_gpt":0.3516074982914908,"score_spread":0.3197820423248843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243745316","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00048961584,0.00012783775,0.00019924669,0.00011108663,0.000025113828,0.00001634146,0.99706763,0.0005610171,0.0014020713],"genre_scores_gemma":[0.001254561,0.000106640975,0.0008398765,0.000065444336,0.000006233987,0.00003873221,0.99704534,0.000038905982,0.0006043193],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992378,0.00007091253,0.00008999714,0.0002539979,0.0002412002,0.0001060292],"domain_scores_gemma":[0.9977387,0.00044293658,0.00031200107,0.0005607997,0.00054536766,0.00040019915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008301408,0.0013728446,0.0008772815,0.0049314247,0.000660181,0.0016256603,0.0013724106,0.0016855338,0.042995326],"category_scores_gemma":[0.004002291,0.00041074576,0.0007280011,0.006983704,0.00030185792,0.0010611553,0.0018059073,0.001066827,0.044801306],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015482356,0.000041045525,0.0045384387,0.0015432994,0.00004364226,0.000113852366,0.00006378626,0.000412952,0.0014978664,0.00089053134,0.97900385,0.0116958385],"study_design_scores_gemma":[0.00008024483,0.000022271883,0.01187053,0.00024883228,0.000035819732,0.00021201826,0.000112433365,0.0010573324,0.00193914,0.0010129968,0.983383,0.000025376557],"about_ca_topic_score_codex":0.01689655,"about_ca_topic_score_gemma":0.03751084,"teacher_disagreement_score":0.042995326,"about_ca_system_score_codex":0.0013168168,"about_ca_system_score_gemma":0.0024676071,"threshold_uncertainty_score":0.1438337},"labels":[],"label_agreement":null},{"id":"W4243805521","doi":"10.1016/j.jvir.2009.04.031","title":"The IR Radlex Project: An Interventional Radiology Lexicon—A Collaborative Project of the Radiological Society of North America and the Society of Interventional Radiology","year":2009,"lang":"en","type":"article","venue":"Journal of Vascular and Interventional Radiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute","funders":"","keywords":"Subspecialty; Medicine; Terminology; Radiology; Variety (cybernetics); Lexicon; Standardization; Radiology information systems; Medical physics; Computer science; Pathology; Artificial intelligence","score_opus":0.01862374785845283,"score_gpt":0.3032958912330144,"score_spread":0.2846721433745616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4243805521","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.031055951,0.0058105323,0.6019734,0.010408389,0.0013272177,0.0036240106,0.22399504,0.05512719,0.06667831],"genre_scores_gemma":[0.060382035,0.0038644185,0.5749986,0.0025086885,0.00044529774,0.0018005946,0.3390687,0.0074022585,0.009529411],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99708325,0.0009542812,0.00066846673,0.00045305685,0.0007337547,0.00010721896],"domain_scores_gemma":[0.9839302,0.0076413047,0.0018689312,0.0014312554,0.004016071,0.0011123341],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004465125,0.0012630179,0.001331337,0.012458281,0.0014936777,0.0051247696,0.0024386446,0.0018560676,0.013691816],"category_scores_gemma":[0.021350741,0.00089206034,0.0013269875,0.007864684,0.0010910223,0.0057610874,0.0037820642,0.0024477188,0.0076495255],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060491636,0.00058297755,0.008448362,0.0066402545,0.0003510814,0.0011116845,0.0017739054,0.004827448,0.024478763,0.0697554,0.44496065,0.4364646],"study_design_scores_gemma":[0.0002594392,0.00014632559,0.010165018,0.0014992215,0.0007131092,0.0020700172,0.0010602282,0.023770306,0.010895482,0.023267109,0.92596465,0.00018912891],"about_ca_topic_score_codex":0.010982746,"about_ca_topic_score_gemma":0.011403782,"teacher_disagreement_score":0.013691816,"about_ca_system_score_codex":0.003045716,"about_ca_system_score_gemma":0.011881599,"threshold_uncertainty_score":0.045803666},"labels":[],"label_agreement":null},{"id":"W4244522263","doi":"10.1002/asi.20463","title":"The reusability of induced knowledge for the automatic semantic markup of taxonomic descriptions","year":2006,"lang":"en","type":"article","venue":"Journal of the American Society for Information Science and Technology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Markup language; Computer science; Natural language processing; Knowledge base; XML; Artificial intelligence; Information retrieval; Reusability; Text corpus; World Wide Web; Programming language; Software","score_opus":0.015963442106323035,"score_gpt":0.2855851794837944,"score_spread":0.26962173737747136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4244522263","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1228889,0.00033950535,0.8614857,0.00051613624,0.00006828222,0.00029470434,0.0013907624,0.008289433,0.0047264933],"genre_scores_gemma":[0.35715497,0.00021828485,0.6371978,0.00012362546,0.000031636904,0.00021550075,0.0034418749,0.00057496305,0.0010412368],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99437535,0.003033476,0.0005580852,0.0008012558,0.0011319062,0.00009998209],"domain_scores_gemma":[0.95078146,0.03122158,0.0027466598,0.0115618585,0.00343634,0.0002521241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007645577,0.00038887953,0.00041833663,0.0036454434,0.000777719,0.001756452,0.0015397356,0.0007446471,0.0015663711],"category_scores_gemma":[0.04151201,0.00058571936,0.00053596316,0.0021007394,0.001227902,0.005339296,0.0022046252,0.0012827079,0.00056769815],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044192473,0.00050847785,0.008655181,0.0010605202,0.00014651258,0.0008218074,0.0024624206,0.021903107,0.075978026,0.04392258,0.0067090634,0.83739036],"study_design_scores_gemma":[0.00011006584,0.000457287,0.013566493,0.0005427751,0.00023173977,0.0015432715,0.0012777416,0.49575004,0.36214092,0.056474168,0.06770004,0.00020552114],"about_ca_topic_score_codex":0.0010924892,"about_ca_topic_score_gemma":0.00191861,"teacher_disagreement_score":0.007645577,"about_ca_system_score_codex":0.000768791,"about_ca_system_score_gemma":0.0010512583,"threshold_uncertainty_score":0.04043418},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"simulation_or_modeling","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"medium"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W4245308851","doi":"10.29173/jchla/jabsc.v32i1.22535","title":"Canadian Virtual Health Library / Bibliothèque virtuelle canadienne de la santé (CVHL/BVCS)","year":2014,"lang":"fr","type":"article","venue":"Journal of the Canadian Health Libraries Association / Journal de l Association de bilbiothèques de la santé du Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Library science; Art; Computer science","score_opus":0.003344606273247575,"score_gpt":0.23456201189876955,"score_spread":0.23121740562552198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4245308851","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037087495,0.01529449,0.006547135,0.08478088,0.0063275956,0.0005416128,0.113122545,0.0044436036,0.76523334],"genre_scores_gemma":[0.10055695,0.029447157,0.041165106,0.019194635,0.002349935,0.0007131538,0.11619013,0.0025484436,0.68783456],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99021775,0.0010983022,0.0005550367,0.0005847075,0.006421893,0.0011223326],"domain_scores_gemma":[0.9621519,0.003932132,0.0013222852,0.0020546163,0.021319943,0.009219153],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0074020573,0.00087418937,0.0010680854,0.01134883,0.008020305,0.015831292,0.0027299097,0.0028201912,0.13827156],"category_scores_gemma":[0.034700222,0.00079198834,0.0008275152,0.017079905,0.0024827856,0.00358757,0.004495581,0.0029525906,0.021181855],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030767387,0.000021293417,0.00094594236,0.00025826035,0.000017836745,0.000040901126,0.00012020386,0.00015495642,0.000099650315,0.024948316,0.9176674,0.055694513],"study_design_scores_gemma":[0.000013902772,0.000003697115,0.0040040747,0.00017081005,0.000011974678,0.000037651425,0.00012514679,0.00024232127,0.00013327728,0.0024745662,0.99275726,0.000025312032],"about_ca_topic_score_codex":0.9420765,"about_ca_topic_score_gemma":0.9362297,"teacher_disagreement_score":0.9841687,"about_ca_system_score_codex":0.06294024,"about_ca_system_score_gemma":0.25872552,"threshold_uncertainty_score":0.4625644},"labels":[],"label_agreement":null},{"id":"W4245802466","doi":"10.32920/14638710.v1","title":"The trainees' perspective on developing an end-of-grant knowledge translation plan","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Manitoba; Queen's University; McMaster University; University of Toronto; Hamilton Health Sciences; Toronto Metropolitan University; University of Calgary","funders":"Canadian Institutes of Health Research; Fondation pour la Recherche Médicale; European Observatory on Health Systems and Policies","keywords":"Knowledge translation; Craft; Plan (archaeology); Perspective (graphical); Process (computing); Library science; Political science; Knowledge management; Computer science; Geography; Artificial intelligence","score_opus":0.09461686856878863,"score_gpt":0.35244418726091675,"score_spread":0.2578273186921281,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4245802466","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049777334,0.0006166805,0.014967612,0.951688,0.0026234288,0.00051941857,0.000076438475,0.00011962684,0.024411136],"genre_scores_gemma":[0.2940301,0.0037886512,0.13450748,0.5114185,0.005016836,0.0041628336,0.00039611477,0.0004510788,0.046228427],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.7107527,0.22442833,0.011174103,0.005909307,0.022044573,0.025690883],"domain_scores_gemma":[0.5971288,0.2111573,0.012722928,0.015974686,0.061932746,0.10108344],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.30587605,0.0009281844,0.0009632165,0.0019825592,0.01779482,0.032972913,0.009641612,0.03614795,0.017724415],"category_scores_gemma":[0.30853927,0.0012246788,0.00218483,0.0024535463,0.024448995,0.02313293,0.03273581,0.049379785,0.0053578657],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003118661,0.0010222035,0.0033233643,0.0014239794,0.0000938253,0.0029669649,0.11081159,0.003876629,0.0017411556,0.5137919,0.2549073,0.10572925],"study_design_scores_gemma":[0.00022523958,0.00062293274,0.0013850009,0.0027131448,0.000066863475,0.0008821263,0.0996865,0.003432832,0.0015744589,0.15351138,0.73562044,0.00027907462],"about_ca_topic_score_codex":0.014480207,"about_ca_topic_score_gemma":0.013733644,"teacher_disagreement_score":0.694124,"about_ca_system_score_codex":0.021511449,"about_ca_system_score_gemma":0.19456777,"threshold_uncertainty_score":0.8559784},"labels":[],"label_agreement":null},{"id":"W4245816202","doi":"10.5596/c07-035","title":"Web 3.0 and health librarians: an introduction","year":2008,"lang":"fr","type":"article","venue":"Journal of the Canadian Health Libraries Association / Journal de l Association de bilbiothèques de la santé du Canada","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Stornoway Diamond (Canada)","funders":"","keywords":"World Wide Web; Social Semantic Web; Semantic Web; Computer science; Web standards; Data Web; Web modeling; Web development; Web intelligence; Semantic Web Stack; Web design; Web page; Theme (computing)","score_opus":0.008525527470714053,"score_gpt":0.25026768486780493,"score_spread":0.24174215739709087,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4245816202","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001531175,0.70852435,0.0048179068,0.1286052,0.015514352,0.00012573398,0.00036521908,0.0002169991,0.14029902],"genre_scores_gemma":[0.014832,0.7838597,0.00995942,0.068445355,0.036954623,0.00026740215,0.00055244396,0.000240277,0.08488873],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9978123,0.0008283122,0.00021938796,0.00016234655,0.00070296763,0.00027470166],"domain_scores_gemma":[0.9943658,0.0035930006,0.0002795067,0.00012247433,0.0009178098,0.0007213404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002669128,0.00096030714,0.00092288427,0.007343101,0.0032088216,0.01028367,0.0013867181,0.009367636,0.020009818],"category_scores_gemma":[0.0043847067,0.0009645577,0.00061308814,0.013200579,0.0047789607,0.017162394,0.0043071005,0.007188096,0.0071272776],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004021624,0.00018840033,0.0011861977,0.0024659142,0.000018339995,0.00033516335,0.0038648136,0.00038105054,0.00037178534,0.19838499,0.5405055,0.25225753],"study_design_scores_gemma":[0.0000045286033,0.000024973227,0.0007503227,0.0014567735,0.0000034458403,0.00026957277,0.0010415063,0.00009600884,0.00003239691,0.018906284,0.97739387,0.000020282267],"about_ca_topic_score_codex":0.009061308,"about_ca_topic_score_gemma":0.009677457,"teacher_disagreement_score":0.020009818,"about_ca_system_score_codex":0.004107502,"about_ca_system_score_gemma":0.0033065104,"threshold_uncertainty_score":0.06693947},"labels":[],"label_agreement":null},{"id":"W4246604680","doi":"10.1038/npre.2009.3570","title":"Open Biomedical Ontologies Applied to Prostate Cancer","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"London Health Sciences Centre; Western University","funders":"London Health Sciences Centre","keywords":"SNOMED CT; Ontology; Computer science; Open Biomedical Ontologies; Controlled vocabulary; Information retrieval; DICOM; World Wide Web; Terminology; Upper ontology; Semantic Web; Artificial intelligence; Suggested Upper Merged Ontology","score_opus":0.017104730573146845,"score_gpt":0.3500854965580411,"score_spread":0.3329807659848943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4246604680","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012503035,0.008201162,0.93924266,0.006885339,0.00056844176,0.00041832202,0.0045465846,0.0027160859,0.024918383],"genre_scores_gemma":[0.11608875,0.008783061,0.85730356,0.0015529853,0.00037555015,0.0004645386,0.009572296,0.0007950382,0.0050641648],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98912257,0.004221231,0.0016449826,0.0012199229,0.0034314278,0.00035993458],"domain_scores_gemma":[0.98244286,0.010517791,0.001154928,0.0032519826,0.0021876567,0.00044486075],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008379539,0.0007874003,0.0010336968,0.014851072,0.0029733498,0.0072152982,0.0019330179,0.0013839396,0.0042604413],"category_scores_gemma":[0.03513603,0.0008387879,0.0021175698,0.021003822,0.0028510618,0.008482199,0.006802517,0.0021929073,0.0012651598],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006981056,0.00009399388,0.0039147628,0.0020327263,0.00016539475,0.0009731828,0.0047520665,0.009500069,0.002659193,0.5985297,0.017302915,0.36000618],"study_design_scores_gemma":[0.00002823215,0.000020407533,0.0034291125,0.0011322133,0.00011288984,0.00078658847,0.001828125,0.029729553,0.005114921,0.53283405,0.4248966,0.00008737072],"about_ca_topic_score_codex":0.01722254,"about_ca_topic_score_gemma":0.015146249,"teacher_disagreement_score":0.01722254,"about_ca_system_score_codex":0.005011658,"about_ca_system_score_gemma":0.006693948,"threshold_uncertainty_score":0.044315755},"labels":[],"label_agreement":null},{"id":"W4246778724","doi":"10.1007/978-3-642-57231-9_7","title":"Systematik","year":2000,"lang":"de","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Philosophy; Gynecology; Medicine","score_opus":0.017457804500662293,"score_gpt":0.24190395865612419,"score_spread":0.2244461541554619,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4246778724","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010074731,0.0065461006,0.31853163,0.005032307,0.002084027,0.00060327904,0.03240303,0.09978516,0.52493966],"genre_scores_gemma":[0.09529243,0.009771492,0.28184414,0.0029876449,0.0007783739,0.0016238975,0.12412623,0.023803767,0.45977205],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99759394,0.0005062055,0.0002447668,0.0007153212,0.0007024242,0.00023724517],"domain_scores_gemma":[0.99917704,0.00017667924,0.000032334432,0.00035778314,0.00019263367,0.00006364093],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014961511,0.002629261,0.0014758629,0.0030166055,0.0017513944,0.0077613113,0.002380897,0.001725354,0.1398719],"category_scores_gemma":[0.003107126,0.0014112407,0.0013155657,0.0034415775,0.0011478596,0.008053291,0.0057394314,0.0056587225,0.19704479],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022403762,0.00020377022,0.00069041486,0.0011956526,0.00010205105,0.00023647951,0.00072518026,0.0024051585,0.004292515,0.1644849,0.2753496,0.55009025],"study_design_scores_gemma":[0.00007748261,0.00004388448,0.0004243246,0.00023475842,0.000058783033,0.00044047396,0.00019715622,0.007243993,0.0049294718,0.11094607,0.8753513,0.000052331103],"about_ca_topic_score_codex":0.0022352492,"about_ca_topic_score_gemma":0.0018092652,"teacher_disagreement_score":0.1398719,"about_ca_system_score_codex":0.0017820022,"about_ca_system_score_gemma":0.0021667217,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4247203180","doi":"10.3410/f.733044204.793548095","title":"Faculty Opinions recommendation of Plain-language medical vocabulary for precision diagnosis.","year":2018,"lang":"en","type":"dataset","venue":"Faculty Opinions – Post-Publication Peer Review of the Biomedical Literature","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Vocabulary; Plain language; Natural language processing; Plain English; Computer science; Artificial intelligence; Information retrieval; Linguistics; Philosophy","score_opus":0.02874782513973628,"score_gpt":0.3765848443751871,"score_spread":0.3478370192354508,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247203180","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00060287176,0.00022167226,0.00070667936,0.00034710544,0.00015572648,0.00009615125,0.9906342,0.0033119707,0.003923661],"genre_scores_gemma":[0.001055237,0.00010460535,0.0016955453,0.000118521864,0.000026211063,0.00007764008,0.9950559,0.00019547757,0.0016709282],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.997335,0.0003480949,0.00049111474,0.0006694678,0.00086959935,0.00028672707],"domain_scores_gemma":[0.99324036,0.0016077346,0.0004950117,0.0016413169,0.0022537706,0.0007617815],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025447782,0.0023850922,0.0014607128,0.007694289,0.00082560873,0.003336658,0.0029550297,0.0025954503,0.09607262],"category_scores_gemma":[0.020651625,0.00076431385,0.0014406658,0.0070983577,0.000565228,0.002693148,0.0029463167,0.0020633747,0.11540987],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009584356,0.00004094078,0.0009166715,0.00054152496,0.000024987221,0.000024940702,0.000019162919,0.00011843253,0.00035281264,0.0004112837,0.9895476,0.007905768],"study_design_scores_gemma":[0.00030739026,0.000036123925,0.0032103357,0.00031996318,0.0000438033,0.00013012446,0.000071994495,0.0016625724,0.0014919954,0.0015593924,0.99113464,0.00003173468],"about_ca_topic_score_codex":0.016158914,"about_ca_topic_score_gemma":0.033699702,"teacher_disagreement_score":0.09607262,"about_ca_system_score_codex":0.001497041,"about_ca_system_score_gemma":0.004807781,"threshold_uncertainty_score":0.32139492},"labels":[],"label_agreement":null},{"id":"W4247290115","doi":"10.1007/978-0-387-39940-9_3963","title":"Visual Discovery","year":2009,"lang":"en","type":"book-chapter","venue":"Encyclopedia of Database Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science","score_opus":0.012913445041344798,"score_gpt":0.26104208070044016,"score_spread":0.24812863565909538,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247290115","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00546313,0.0046057375,0.5008157,0.0023097387,0.00097181654,0.00081162923,0.06604607,0.09757104,0.32140526],"genre_scores_gemma":[0.06894087,0.0061479076,0.4564879,0.002027096,0.00042808632,0.0009663409,0.112848006,0.018902624,0.3332511],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994574,0.000050841718,0.000034300363,0.00016244513,0.00022794893,0.00006704448],"domain_scores_gemma":[0.9990852,0.00021560794,0.000036212634,0.00021140526,0.00036899737,0.000082609404],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00061087817,0.0011513946,0.00057693524,0.0048827333,0.00075672695,0.005066085,0.0016494807,0.000936895,0.1379918],"category_scores_gemma":[0.0024613335,0.00046949938,0.0010246343,0.0030524514,0.00038786247,0.0026456874,0.0021160448,0.00094880385,0.070726454],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016492275,0.000042857722,0.0004095896,0.0007132852,0.000046778987,0.00013744009,0.0001906226,0.0011452757,0.0066015297,0.026554061,0.444425,0.5195686],"study_design_scores_gemma":[0.000034281267,0.00002391562,0.0007707348,0.00022971861,0.000032207958,0.00037559643,0.00020686026,0.011341267,0.009500904,0.031994242,0.9454502,0.00004017867],"about_ca_topic_score_codex":0.007316116,"about_ca_topic_score_gemma":0.010365138,"teacher_disagreement_score":0.1379918,"about_ca_system_score_codex":0.0008501585,"about_ca_system_score_gemma":0.0011778229,"threshold_uncertainty_score":0.46162856},"labels":[],"label_agreement":null},{"id":"W4247439555","doi":"10.7287/peerj.126v0.1/reviews/3","title":"Peer Review #3 of \"Kinome Render: a stand-alone and web-accessible tool to annotate the human protein kinome tree (v0.1)\"","year":2013,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Kinome; Computer science; Tree (set theory); Computational biology; World Wide Web; Chemistry; Biology; Kinase; Biochemistry; Mathematics","score_opus":0.045471007919794526,"score_gpt":0.3400595925082648,"score_spread":0.29458858458847026,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4247439555","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071687344,0.0045208912,0.0410165,0.079383895,0.317912,0.011589059,0.036632072,0.03729218,0.46448466],"genre_scores_gemma":[0.016389405,0.0031141096,0.018016873,0.0049351,0.024669066,0.0020728563,0.024861641,0.015140001,0.8908009],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9877033,0.0016395492,0.0009320917,0.00063521834,0.0083812,0.00070863456],"domain_scores_gemma":[0.821946,0.015017211,0.003032046,0.008766448,0.14038421,0.010854028],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.010358254,0.0011110689,0.0015606186,0.0069626183,0.0045716353,0.0068320134,0.0026075444,0.002535631,0.47210893],"category_scores_gemma":[0.09985665,0.00070399186,0.0014035485,0.0030652885,0.0015881002,0.004075075,0.0051575946,0.0018354546,0.37714556],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020127845,0.000011333711,0.00012502687,0.0002682628,0.0000061034057,0.00009369789,0.000071581744,0.00003141484,0.00043602104,0.0003631652,0.9707473,0.027825883],"study_design_scores_gemma":[0.000016193904,0.000013003818,0.0005776772,0.0001543131,0.0000066437533,0.000108292064,0.00009614534,0.00022918369,0.0006107953,0.0006867713,0.9974794,0.000021496753],"about_ca_topic_score_codex":0.004764554,"about_ca_topic_score_gemma":0.012250428,"teacher_disagreement_score":0.47210893,"about_ca_system_score_codex":0.0021334917,"about_ca_system_score_gemma":0.008586577,"threshold_uncertainty_score":0.75297254},"labels":[],"label_agreement":null},{"id":"W4248372822","doi":"10.1038/npre.2009.3589.1","title":"Open Biomedical Ontologies Applied to Prostate Cancer","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"London Health Sciences Centre","funders":"London Health Sciences Centre","keywords":"SNOMED CT; Ontology; Open Biomedical Ontologies; Computer science; Controlled vocabulary; Information retrieval; DICOM; Section (typography); World Wide Web; Upper ontology; Terminology; Semantic Web; Artificial intelligence; Suggested Upper Merged Ontology","score_opus":0.017104730573146845,"score_gpt":0.3500854965580411,"score_spread":0.3329807659848943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4248372822","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012503035,0.008201162,0.93924266,0.006885339,0.00056844176,0.00041832202,0.0045465846,0.0027160859,0.024918383],"genre_scores_gemma":[0.11608875,0.008783061,0.85730356,0.0015529853,0.00037555015,0.0004645386,0.009572296,0.0007950382,0.0050641648],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98912257,0.004221231,0.0016449826,0.0012199229,0.0034314278,0.00035993458],"domain_scores_gemma":[0.98244286,0.010517791,0.001154928,0.0032519826,0.0021876567,0.00044486075],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.008379539,0.0007874003,0.0010336968,0.014851072,0.0029733498,0.0072152982,0.0019330179,0.0013839396,0.0042604413],"category_scores_gemma":[0.03513603,0.0008387879,0.0021175698,0.021003822,0.0028510618,0.008482199,0.006802517,0.0021929073,0.0012651598],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006981056,0.00009399388,0.0039147628,0.0020327263,0.00016539475,0.0009731828,0.0047520665,0.009500069,0.002659193,0.5985297,0.017302915,0.36000618],"study_design_scores_gemma":[0.00002823215,0.000020407533,0.0034291125,0.0011322133,0.00011288984,0.00078658847,0.001828125,0.029729553,0.005114921,0.53283405,0.4248966,0.00008737072],"about_ca_topic_score_codex":0.01722254,"about_ca_topic_score_gemma":0.015146249,"teacher_disagreement_score":0.99806696,"about_ca_system_score_codex":0.005011658,"about_ca_system_score_gemma":0.006693948,"threshold_uncertainty_score":0.044315755},"labels":[],"label_agreement":null},{"id":"W4249267637","doi":"10.1515/iupac.69.0630","title":"Properties and Units in the Clinical Laboratory Sciences: Part XI. Coding Systems - Structure and Guidelines","year":2016,"lang":"en","type":"dataset","venue":"IUPAC Standards Online","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"","keywords":"Syntax; Coding (social sciences); Computer science; Natural language processing; Information retrieval; Mathematics; Statistics","score_opus":0.1021377497889674,"score_gpt":0.44515666500251116,"score_spread":0.3430189152135438,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249267637","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020285188,0.00048774743,0.0021055695,0.00029900146,0.000053173226,0.00009259869,0.99212873,0.0005476753,0.0022569145],"genre_scores_gemma":[0.0037048324,0.0004791943,0.0039353264,0.0001053943,0.000013804659,0.00032975493,0.99043274,0.00013550473,0.0008634234],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9963451,0.00062831945,0.0015566427,0.000530227,0.0007116676,0.00022797342],"domain_scores_gemma":[0.9909428,0.003785446,0.0018139306,0.0016341052,0.0014225258,0.00040125693],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003027118,0.0009968195,0.0007956321,0.006769578,0.00050032703,0.0022569753,0.0012648501,0.000907324,0.014118398],"category_scores_gemma":[0.015210569,0.00062946725,0.0009439492,0.0136755975,0.00064189767,0.0018667164,0.001694025,0.0019187827,0.012069427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058988255,0.000091193615,0.021799313,0.0033752215,0.000096723204,0.0001400609,0.00040416842,0.002266281,0.0012555283,0.01061496,0.90294147,0.056425262],"study_design_scores_gemma":[0.00023720639,0.0000367618,0.029445868,0.0010704728,0.00006406378,0.0002733848,0.00027262353,0.0011894925,0.0019061769,0.0070001325,0.95845556,0.000048196092],"about_ca_topic_score_codex":0.022322888,"about_ca_topic_score_gemma":0.023009453,"teacher_disagreement_score":0.022322888,"about_ca_system_score_codex":0.0029663553,"about_ca_system_score_gemma":0.0059208022,"threshold_uncertainty_score":0.04723078},"labels":[],"label_agreement":null},{"id":"W4249681836","doi":"10.2196/preprints.26123","title":"Examination of a Canada-Wide Collaboration Platform for Order Sets: Retrospective Analysis (Preprint)","year":2020,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Executable; Order (exchange); Preprint; Knowledge translation; Computer science; Upload; Knowledge sharing; Descriptive statistics; Information retrieval; World Wide Web; Knowledge management; Business; Statistics; Mathematics","score_opus":0.017312643514052532,"score_gpt":0.27702229606450307,"score_spread":0.25970965255045053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4249681836","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9153929,0.0030575707,0.0019053576,0.0025851438,0.00010574982,0.0023670713,0.06281113,0.00015296951,0.0116221355],"genre_scores_gemma":[0.9754759,0.002045576,0.0019302671,0.00081724353,0.000059807895,0.0010279587,0.01581271,0.000095580486,0.0027349193],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9811143,0.0022266684,0.002749157,0.0023270312,0.008854622,0.0027281295],"domain_scores_gemma":[0.86368245,0.022924846,0.035420734,0.007115251,0.063996285,0.0068603484],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012526234,0.0006325847,0.0007955692,0.016510027,0.004730697,0.0042037363,0.0023307763,0.00090525666,0.0052640527],"category_scores_gemma":[0.0660917,0.0010734797,0.0013061826,0.034572575,0.0020164684,0.0020358805,0.0032782417,0.0014692667,0.001108558],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014835458,0.000045472218,0.97670776,0.0003173584,0.00008589569,0.00015003751,0.0042161243,0.00010559749,0.000067759545,0.0005949302,0.008982938,0.008577783],"study_design_scores_gemma":[0.00002270244,0.00009270349,0.97361135,0.00055716303,0.00012008719,0.00023277113,0.0149371205,0.0005697422,0.00022757205,0.000110334906,0.009456202,0.00006228687],"about_ca_topic_score_codex":0.94456124,"about_ca_topic_score_gemma":0.9194505,"teacher_disagreement_score":0.9504247,"about_ca_system_score_codex":0.049575347,"about_ca_system_score_gemma":0.11089275,"threshold_uncertainty_score":0.35969597},"labels":[],"label_agreement":null},{"id":"W4250369710","doi":"10.3115/1654415.1654439","title":"Biomedical term recognition with the perceptron HMM algorithm","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Computer science; Artificial intelligence; Hidden Markov model; Term (time); Perceptron; Machine learning; Multilayer perceptron; Pattern recognition (psychology); Identification (biology); Set (abstract data type); Feature (linguistics); Algorithm; Artificial neural network","score_opus":0.00856032494501333,"score_gpt":0.22995595804100483,"score_spread":0.2213956330959915,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250369710","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014964656,0.00095098914,0.9775964,0.00030819586,0.00016762054,0.00008218982,0.0004290045,0.004415228,0.0010857073],"genre_scores_gemma":[0.2786254,0.0009210514,0.7137911,0.00031600194,0.0002633714,0.00031773135,0.0015804017,0.00016958897,0.0040154704],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99859244,0.00038540026,0.0001903497,0.00035293575,0.000373743,0.00010509032],"domain_scores_gemma":[0.996977,0.0019689142,0.00023687333,0.0002462992,0.00049396895,0.000077023346],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022558244,0.0006972041,0.0010766576,0.0020037815,0.0004022539,0.0013203947,0.0015181914,0.0012849353,0.0019828593],"category_scores_gemma":[0.0072942753,0.00042538065,0.0008745419,0.0020827735,0.00050245045,0.0020768603,0.0009014008,0.0015552887,0.0019701396],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006486138,0.00027853693,0.0048538987,0.0004660792,0.00022648694,0.00035549494,0.00030508815,0.08346215,0.035630062,0.007942796,0.008635883,0.8571949],"study_design_scores_gemma":[0.000048867794,0.00009870643,0.0017089703,0.000040735784,0.00006917803,0.00018125531,0.00003849698,0.9673214,0.013143592,0.013316241,0.003977761,0.00005480883],"about_ca_topic_score_codex":0.004257913,"about_ca_topic_score_gemma":0.0033233012,"teacher_disagreement_score":0.004257913,"about_ca_system_score_codex":0.0006828018,"about_ca_system_score_gemma":0.0009451871,"threshold_uncertainty_score":0.011930108},"labels":[],"label_agreement":null},{"id":"W4250850010","doi":"10.4018/978-1-60960-561-2.ch710","title":"Social Cognitive Ontology and User Driven Healthcare","year":2011,"lang":"en","type":"book-chapter","venue":"IGI Global eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"NOSM University","funders":"","keywords":"Ontology; Health care; Premise; Knowledge management; Meaning (existential); Cognition; Key (lock); Process (computing); Feeling; Computer science; Psychology; Epistemology; Social psychology; Psychotherapist; Political science","score_opus":0.03469957915362512,"score_gpt":0.2963592641442488,"score_spread":0.26165968499062364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4250850010","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02439514,0.013599587,0.4962888,0.0698317,0.0012537439,0.0003122122,0.00028663006,0.0006545132,0.39337763],"genre_scores_gemma":[0.8241762,0.007683561,0.13399202,0.0055169985,0.0008794234,0.00084213354,0.00034214288,0.0002440286,0.026323529],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9930074,0.004458166,0.0003236417,0.0005101738,0.001226753,0.00047385518],"domain_scores_gemma":[0.9935314,0.0046439483,0.00030700606,0.0007111171,0.00046320164,0.00034329464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073232222,0.0008106067,0.00064451643,0.0027176957,0.0029256092,0.009489728,0.0019946517,0.0040605557,0.005135034],"category_scores_gemma":[0.0068443837,0.0005733512,0.0012499071,0.0024174913,0.02858799,0.010539076,0.0065020267,0.004167202,0.000568279],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000035163164,0.0000094961015,0.00008290434,0.000036249843,0.0000050062554,0.0000577535,0.0036206918,0.00047441333,0.000050393723,0.9917292,0.0008057177,0.0031245938],"study_design_scores_gemma":[0.000007486182,0.000011665612,0.00014400348,0.00008406068,0.0000056810504,0.000118536074,0.0025643539,0.0033879557,0.00011257856,0.9409702,0.052578367,0.000015129357],"about_ca_topic_score_codex":0.0067845094,"about_ca_topic_score_gemma":0.003264579,"teacher_disagreement_score":0.009489728,"about_ca_system_score_codex":0.008053455,"about_ca_system_score_gemma":0.0042354995,"threshold_uncertainty_score":0.05843216},"labels":[],"label_agreement":null},{"id":"W4252085617","doi":"10.1038/npre.2009.3542.1","title":"An OWL-DL Ontology for Classification of Lipids","year":2009,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Ontology; Computer science; Web Ontology Language; Computational biology; Axiom; Information retrieval; Semantic Web; Biology; Mathematics; Epistemology","score_opus":0.021081855611997925,"score_gpt":0.34797951611155203,"score_spread":0.3268976604995541,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252085617","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008046641,0.0005323206,0.95714486,0.0019302465,0.00036666164,0.0006088538,0.01557185,0.0051116673,0.010686915],"genre_scores_gemma":[0.05097483,0.00092378474,0.917139,0.0008683362,0.000123644,0.0005168427,0.024630316,0.0005478139,0.004275426],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99748397,0.0004849977,0.00053383695,0.00031478994,0.0010454657,0.00013697222],"domain_scores_gemma":[0.9971431,0.00085252867,0.00024627798,0.00066000904,0.00085030176,0.00024784196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035284616,0.00063334935,0.0006557963,0.0034925535,0.0014771332,0.0032430596,0.0021609026,0.0010790388,0.002960928],"category_scores_gemma":[0.0052904054,0.00047959838,0.0012360592,0.004210039,0.0010676775,0.004137393,0.001944321,0.0021519596,0.0015729481],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020632736,0.00046516082,0.003860528,0.0014213255,0.00013270104,0.0019301241,0.0011939667,0.012224575,0.017085535,0.61809206,0.08870506,0.25468266],"study_design_scores_gemma":[0.0001224904,0.00004039027,0.0023306122,0.00050998543,0.00011422216,0.0015754611,0.00050426146,0.12506355,0.011823836,0.2730545,0.5847713,0.00008938866],"about_ca_topic_score_codex":0.014375173,"about_ca_topic_score_gemma":0.009183271,"teacher_disagreement_score":0.014375173,"about_ca_system_score_codex":0.0025231086,"about_ca_system_score_gemma":0.004833039,"threshold_uncertainty_score":0.02858299},"labels":[],"label_agreement":null},{"id":"W4252356357","doi":"10.5489/cuaj.7433","title":"Authorship error","year":2021,"lang":"en","type":"erratum","venue":"Canadian Urological Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.019526485506873758,"score_gpt":0.25520537568067897,"score_spread":0.23567889017380522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252356357","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00048624707,0.00089926337,0.0017996541,0.19074427,0.75346273,0.00010800863,0.004876044,0.0007860356,0.04683782],"genre_scores_gemma":[0.02222011,0.0020544087,0.0053837695,0.17351562,0.08604622,0.00052839017,0.0049395203,0.0020483593,0.7032637],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.978744,0.0029561606,0.0025899007,0.0030549015,0.010861638,0.0017935049],"domain_scores_gemma":[0.89132303,0.033639867,0.0045067854,0.008267528,0.058474887,0.0037879595],"candidate_categories":["research_integrity"],"consensus_categories":[],"category_scores_codex":[0.0103853,0.001062608,0.0018227171,0.0036166369,0.004954355,0.005756351,0.0024702332,0.008714831,0.14133842],"category_scores_gemma":[0.15138656,0.00071832165,0.001098222,0.0036067101,0.0029696333,0.0024770824,0.0040975157,0.009073981,0.0829599],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020115995,0.0000050118397,0.00007232479,0.000055284687,0.0000045239303,0.00013904295,0.00008521244,0.000016970089,0.00002169119,0.0025163603,0.99142176,0.005641644],"study_design_scores_gemma":[0.00002217424,0.0000062180443,0.00025648237,0.0003250092,0.000013373715,0.00027666395,0.00018561573,0.00011316678,0.00021171031,0.0028501276,0.99572265,0.000016700269],"about_ca_topic_score_codex":0.018372824,"about_ca_topic_score_gemma":0.029887307,"teacher_disagreement_score":0.99128515,"about_ca_system_score_codex":0.0074345316,"about_ca_system_score_gemma":0.010143977,"threshold_uncertainty_score":0.47282416},"labels":[],"label_agreement":null},{"id":"W4252790521","doi":"10.1038/npre.2007.945.2","title":"Bridging the gap between social tagging and semantic annotation: E.D. the Entity Describer","year":2007,"lang":"en","type":"preprint","venue":"Nature Precedings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Providence Health Care","funders":"Genome Prairie; Genome Alberta","keywords":"Computer science; Bridging (networking); Annotation; Information retrieval; RDF; Semantic Web; Semantic Web Stack; World Wide Web; Artificial intelligence","score_opus":0.02961210909848352,"score_gpt":0.3225281930507134,"score_spread":0.2929160839522299,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252790521","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013318755,0.0008781291,0.9519235,0.0105493795,0.0002003792,0.00015747291,0.0007443514,0.0033019406,0.018926065],"genre_scores_gemma":[0.23489828,0.0018693174,0.74156934,0.0021076503,0.0003487437,0.00025828936,0.0025336668,0.0019304996,0.014484196],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9824806,0.010127111,0.001198079,0.0020037757,0.003760822,0.0004296724],"domain_scores_gemma":[0.9319052,0.03467353,0.0019601774,0.026408434,0.0041420804,0.0009105905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027081467,0.0007065624,0.0011213372,0.004713783,0.0026693754,0.008005675,0.0029525133,0.00291575,0.0065598697],"category_scores_gemma":[0.048340566,0.0009734538,0.0011431779,0.004802657,0.004886254,0.022941563,0.009905927,0.002859911,0.003215576],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029355593,0.00015685278,0.0025680666,0.0003124389,0.000068207955,0.0005416507,0.00338904,0.004518724,0.0052407742,0.7530576,0.023690019,0.20616305],"study_design_scores_gemma":[0.00009595103,0.00005581402,0.0010178942,0.00034317438,0.0001386095,0.000990889,0.0020136938,0.11951813,0.027870288,0.4954214,0.35238212,0.00015211565],"about_ca_topic_score_codex":0.0096648205,"about_ca_topic_score_gemma":0.005286058,"teacher_disagreement_score":0.027081467,"about_ca_system_score_codex":0.0031275707,"about_ca_system_score_gemma":0.0032747877,"threshold_uncertainty_score":0.14322221},"labels":[],"label_agreement":null},{"id":"W4252854485","doi":"10.33137/js.v1i0.27065","title":"Authority Delegation","year":2016,"lang":"en","type":"article","venue":"Scientonomy Journal for the Science of Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Delegation; Primary authority; Delegated authority; Political science; Public relations; Law and economics; Public administration; Law; Sociology; Traditional authority; Legal research","score_opus":0.030221197333186532,"score_gpt":0.3298692128404677,"score_spread":0.2996480155072812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4252854485","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011703163,0.0022659304,0.39110377,0.011982199,0.0018759352,0.0003626761,0.00046619854,0.0010402976,0.57919985],"genre_scores_gemma":[0.7265445,0.0027470668,0.09589076,0.006576047,0.003593981,0.0009522957,0.00085054437,0.00091360597,0.16193123],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9771947,0.007425403,0.0019116802,0.006154185,0.0052333395,0.002080594],"domain_scores_gemma":[0.97632766,0.008425926,0.0018684332,0.007353805,0.004397695,0.0016264615],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.013006079,0.001307444,0.0012998037,0.003382023,0.006201536,0.009975929,0.0028189833,0.0041025956,0.025475001],"category_scores_gemma":[0.030013382,0.00092815934,0.0022687141,0.0024190603,0.015595026,0.020529563,0.010217388,0.0052425107,0.007017947],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013941144,0.000008366151,0.00017623702,0.000048206355,0.000010088347,0.0000639408,0.0007628808,0.00028957153,0.00013577563,0.9887833,0.002941761,0.0067659663],"study_design_scores_gemma":[0.00002990773,0.000021055183,0.00019838408,0.00009151206,0.000024964835,0.00021233753,0.0003479865,0.0018085976,0.0005802016,0.82501084,0.17163661,0.000037556616],"about_ca_topic_score_codex":0.0033723759,"about_ca_topic_score_gemma":0.002044816,"teacher_disagreement_score":0.996618,"about_ca_system_score_codex":0.00534681,"about_ca_system_score_gemma":0.004527078,"threshold_uncertainty_score":0.08522236},"labels":[],"label_agreement":null},{"id":"W4253244973","doi":"10.1007/978-1-4939-7131-2_100131","title":"Collaborative Tagging","year":2018,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Natural language processing; Artificial intelligence; World Wide Web","score_opus":0.012140756733363763,"score_gpt":0.2601407953641144,"score_spread":0.24800003863075062,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253244973","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003050541,0.0010690856,0.8964228,0.0007510689,0.0007002791,0.0002772932,0.0017152374,0.005889579,0.09012409],"genre_scores_gemma":[0.07658721,0.0021566513,0.7162923,0.0010354213,0.0005503099,0.00042445684,0.012742538,0.0021514315,0.18805975],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9965364,0.0008156487,0.00023230519,0.0012248817,0.0009902726,0.00020052333],"domain_scores_gemma":[0.99386084,0.0016848927,0.0001683236,0.0031938113,0.0009025959,0.0001894973],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033276256,0.00132916,0.0011071949,0.0041671623,0.0023784507,0.0053458544,0.0032855244,0.0023231625,0.03961033],"category_scores_gemma":[0.008671148,0.0008038948,0.0015809031,0.0055503044,0.0010545539,0.008678783,0.005991477,0.001826794,0.039874494],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014333014,0.00015976539,0.0007002986,0.00039623887,0.00010419487,0.00022969668,0.00050120585,0.0037458432,0.0074137533,0.109060794,0.08692547,0.79061943],"study_design_scores_gemma":[0.000034797922,0.000070158916,0.0009315072,0.00028436375,0.00017698608,0.0012493415,0.00048227017,0.069878004,0.026268858,0.27421904,0.6262827,0.00012201034],"about_ca_topic_score_codex":0.0032189453,"about_ca_topic_score_gemma":0.00495604,"teacher_disagreement_score":0.03961033,"about_ca_system_score_codex":0.0010715796,"about_ca_system_score_gemma":0.0020456286,"threshold_uncertainty_score":0.13250977},"labels":[],"label_agreement":null},{"id":"W4253950338","doi":"10.1515/iupac.76.0225","title":"Epithelium","year":2016,"lang":"en","type":"dataset","venue":"IUPAC Standards Online","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Glossary; Toxicokinetics; Relation (database); Hazard; Computer science; Toxicology; Medicine; Chemistry; Pharmacology; Data mining; Biology; Linguistics; Philosophy","score_opus":0.012920860189941038,"score_gpt":0.39421652880590224,"score_spread":0.3812956686159612,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4253950338","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00045026102,0.00022066139,0.0004391708,0.00017755564,0.00008312054,0.000056382985,0.9927757,0.0008978355,0.0048993067],"genre_scores_gemma":[0.0005469737,0.00010689038,0.0007665496,0.00015174245,0.0000111278705,0.0001273702,0.9958883,0.000110172754,0.0022909436],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99766773,0.00034057436,0.0003691467,0.0008277145,0.0005721007,0.00022277908],"domain_scores_gemma":[0.99743515,0.0005419749,0.00023277782,0.00081567856,0.0007688738,0.00020561545],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012403771,0.0017040434,0.0010947541,0.0033799869,0.001084481,0.0029821072,0.0023304792,0.0015004434,0.10896509],"category_scores_gemma":[0.0067143817,0.0005169748,0.001498047,0.004789222,0.00042576564,0.0025972137,0.0024106577,0.001612051,0.15631108],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000096357086,0.000034757126,0.0014886763,0.0006692482,0.000023517716,0.000039416587,0.00004810708,0.00021995607,0.00026576844,0.0013638355,0.9835827,0.012167784],"study_design_scores_gemma":[0.000051329294,0.000014092918,0.0024184608,0.00022221389,0.000015055896,0.000090975765,0.00009576131,0.00028860962,0.0003375596,0.0013016585,0.99514705,0.000017307908],"about_ca_topic_score_codex":0.0111605255,"about_ca_topic_score_gemma":0.023931053,"teacher_disagreement_score":0.8910349,"about_ca_system_score_codex":0.0015085578,"about_ca_system_score_gemma":0.002218616,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4255154157","doi":"10.12688/f1000research.11389.1","title":"PubRunner: A light-weight framework for updating text mining results","year":2017,"lang":"en","type":"preprint","venue":"F1000Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Upload; Workflow; Computer science; Field (mathematics); Data science; Biomedical text mining; Domain (mathematical analysis); World Wide Web; Information retrieval; Text mining; Data mining; Database","score_opus":0.07276717202123249,"score_gpt":0.3943506542953737,"score_spread":0.3215834822741412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255154157","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009099718,0.00039734525,0.5391606,0.001406214,0.0006581776,0.0008133154,0.016206257,0.4369797,0.0034683708],"genre_scores_gemma":[0.010724932,0.00058973586,0.87757814,0.0009401924,0.00045909735,0.0013700406,0.048988305,0.053097323,0.0062522506],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9789892,0.00553646,0.00405598,0.004133241,0.0064434493,0.00084170396],"domain_scores_gemma":[0.9085881,0.04301635,0.005613769,0.025424402,0.01324763,0.004109605],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.045587145,0.0061196918,0.0036454326,0.021837007,0.004623567,0.016145585,0.011229757,0.0045856046,0.04309322],"category_scores_gemma":[0.13453361,0.005451878,0.005849256,0.014878305,0.0033331048,0.027778164,0.015155483,0.006063589,0.054902792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011173503,0.00034229807,0.0034213818,0.002253062,0.00045813763,0.00081190496,0.0018590823,0.006001483,0.006028613,0.040050123,0.54356605,0.39409047],"study_design_scores_gemma":[0.00058959087,0.00019840951,0.0017625145,0.0010123572,0.00025121745,0.0007275613,0.0005261495,0.11214424,0.019062405,0.14053014,0.7225893,0.0006061727],"about_ca_topic_score_codex":0.009945871,"about_ca_topic_score_gemma":0.016751284,"teacher_disagreement_score":0.045587145,"about_ca_system_score_codex":0.0031375354,"about_ca_system_score_gemma":0.008575344,"threshold_uncertainty_score":0.24109071},"labels":[],"label_agreement":null},{"id":"W4255627255","doi":"10.18653/v1/2020.clinicalnlp-1","title":"Proceedings of the 3rd Clinical Natural Language Processing Workshop","year":2020,"lang":"en","type":"paratext","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":48,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Natural language processing; Natural (archaeology); Natural language; Programming language; Artificial intelligence; History; Archaeology","score_opus":0.02644176150270859,"score_gpt":0.34536779823057906,"score_spread":0.31892603672787045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4255627255","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026680987,0.014089641,0.60315895,0.05840622,0.021227088,0.0014815344,0.050447155,0.03170846,0.19279997],"genre_scores_gemma":[0.045341585,0.007953288,0.32106605,0.0067858184,0.003151575,0.0010997556,0.12791005,0.008574982,0.47811684],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99777454,0.0010732805,0.00019275717,0.00028265812,0.00053681247,0.0001398712],"domain_scores_gemma":[0.9949173,0.002238092,0.00008274403,0.0006526148,0.0014019164,0.00070739974],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006443244,0.0011282332,0.0012303038,0.0015120547,0.0008369408,0.0054368046,0.0015068838,0.001478844,0.07114818],"category_scores_gemma":[0.007257235,0.0005375504,0.0011643254,0.0013480848,0.0007052449,0.0035568802,0.0027806868,0.002038201,0.0404227],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006161761,0.0003450462,0.00058548985,0.00029883638,0.00007027486,0.00027733346,0.0003624718,0.001349926,0.0034546868,0.006623886,0.7230608,0.26295504],"study_design_scores_gemma":[0.00014556538,0.00007866545,0.0012099422,0.0001633238,0.00007089752,0.00030311852,0.00024557376,0.01604206,0.0057126237,0.014149408,0.9618442,0.000034546036],"about_ca_topic_score_codex":0.006082323,"about_ca_topic_score_gemma":0.010504399,"teacher_disagreement_score":0.92885184,"about_ca_system_score_codex":0.0013747024,"about_ca_system_score_gemma":0.004132166,"threshold_uncertainty_score":0.2380144},"labels":[],"label_agreement":null},{"id":"W4256721440","doi":"10.1007/978-1-4939-7131-2_101029","title":"Scientific Collaboration Network","year":2018,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science","score_opus":0.017506341661761668,"score_gpt":0.26077396170834016,"score_spread":0.24326762004657848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4256721440","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017027795,0.0044249115,0.41869342,0.0101914415,0.0012669752,0.00065277104,0.024193766,0.00324963,0.52029926],"genre_scores_gemma":[0.29785272,0.014123018,0.31631228,0.0018853232,0.0014106528,0.0017471996,0.061737336,0.0012027896,0.3037287],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9980354,0.0004996042,0.00017661475,0.0005446336,0.00062367984,0.00012020636],"domain_scores_gemma":[0.9976908,0.00084921205,0.0002678012,0.00042213927,0.0004456456,0.00032441263],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0016039662,0.0006624632,0.00046479548,0.0064036995,0.0017945615,0.0042603198,0.0011281483,0.0012144204,0.035174627],"category_scores_gemma":[0.005423038,0.00034676434,0.000732584,0.008303279,0.00077921647,0.0072887638,0.0047341264,0.0010939403,0.011857176],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006234397,0.00003666894,0.0012579946,0.00049893465,0.000068668945,0.00044018132,0.0004628036,0.003020032,0.0019722902,0.7727347,0.07638366,0.14306168],"study_design_scores_gemma":[0.0000107550195,0.000014605709,0.0006117746,0.00016378607,0.00004192553,0.0005789236,0.00026529207,0.0069272956,0.0010422807,0.35495058,0.6353757,0.000017083164],"about_ca_topic_score_codex":0.0015441446,"about_ca_topic_score_gemma":0.0014790024,"teacher_disagreement_score":0.9935963,"about_ca_system_score_codex":0.0011756155,"about_ca_system_score_gemma":0.0021896763,"threshold_uncertainty_score":0.11767089},"labels":[],"label_agreement":null},{"id":"W4280500719","doi":"10.1016/j.isci.2022.104390","title":"A graph-embedded topic model enables characterization of diverse pain phenotypes among UK biobank individuals","year":2022,"lang":"en","type":"article","venue":"iScience","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Medical Research Council; McGill University","keywords":"Biobank; Autoencoder; Phenotype; Chronic pain; Graph; Computer science; Medicine; Data science; Data mining; Machine learning; Artificial intelligence; Bioinformatics; Theoretical computer science; Artificial neural network; Psychiatry; Gene; Biology","score_opus":0.017105600673552074,"score_gpt":0.2458877478001447,"score_spread":0.22878214712659262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4280500719","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30901828,0.0019220855,0.6682275,0.0020795418,0.00013993756,0.00019610292,0.013961507,0.0018278898,0.0026271057],"genre_scores_gemma":[0.86345404,0.0008773869,0.10850448,0.0004129476,0.00017361957,0.00030018575,0.02268421,0.00020812404,0.0033850558],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993143,0.00024260151,0.000041830306,0.00029059532,0.000053080057,0.000057610876],"domain_scores_gemma":[0.9970222,0.0022584773,0.00022511433,0.00026638468,0.00016942763,0.000058304275],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014044751,0.00055262167,0.00054737757,0.0018966313,0.00040047878,0.0007796312,0.0009332556,0.0012807515,0.0019257151],"category_scores_gemma":[0.006109561,0.00037063242,0.0012506412,0.0018179234,0.0004817473,0.0012356366,0.00091644336,0.0012427961,0.0007865244],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010370174,0.00054145185,0.14436959,0.0007180529,0.001131431,0.000999783,0.002089342,0.4578522,0.012811809,0.039707325,0.032815423,0.30592662],"study_design_scores_gemma":[0.000055626253,0.00006830391,0.020127501,0.00005587177,0.00017865229,0.00028137007,0.00013757811,0.9405163,0.0010552903,0.030708047,0.0067830496,0.000032379612],"about_ca_topic_score_codex":0.014371706,"about_ca_topic_score_gemma":0.025195424,"teacher_disagreement_score":0.014371706,"about_ca_system_score_codex":0.00067848625,"about_ca_system_score_gemma":0.0006065556,"threshold_uncertainty_score":0.028576076},"labels":[],"label_agreement":null},{"id":"W4280505624","doi":"10.1186/s12874-022-01583-z","title":"Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"BMC Medical Research Methodology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia Hospital; Prevention of Organ Failure; University of British Columbia","funders":"University of British Columbia","keywords":"Computer science; Data extraction; Pipeline (software); Artificial intelligence; Medicine; Test (biology); Breast cancer; Medical physics; Natural language processing; Machine learning; Cancer; MEDLINE","score_opus":0.44980429708784625,"score_gpt":0.6128187548373685,"score_spread":0.16301445774952222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4280505624","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04603071,0.0015030737,0.74118924,0.0021209074,0.00035781864,0.0033366748,0.05135918,0.14986975,0.004232698],"genre_scores_gemma":[0.07702602,0.0005272246,0.872169,0.00052787235,0.00022544697,0.0019863385,0.044630434,0.00078788935,0.0021196909],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949051,0.0010444224,0.0011290212,0.0015267157,0.0012419239,0.00015291061],"domain_scores_gemma":[0.9858915,0.0070032547,0.0018812227,0.0013621392,0.0034548454,0.00040692906],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004661462,0.0014182572,0.001090727,0.0056707067,0.00074494135,0.0022515506,0.0019175502,0.0010155835,0.004556406],"category_scores_gemma":[0.017195893,0.00062980637,0.00125461,0.0033262179,0.00040069607,0.002518205,0.002192393,0.0010569637,0.0042078146],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007800992,0.00065756874,0.014021879,0.0025367339,0.00025904409,0.001851643,0.0011327937,0.005209648,0.039726347,0.0024932874,0.09545905,0.8358719],"study_design_scores_gemma":[0.00123828,0.0010202808,0.054583527,0.0010970292,0.00072661415,0.005084754,0.0013485315,0.5814076,0.1336571,0.02303428,0.19620219,0.00059979386],"about_ca_topic_score_codex":0.0050846157,"about_ca_topic_score_gemma":0.00606077,"teacher_disagreement_score":0.0056707067,"about_ca_system_score_codex":0.0013401022,"about_ca_system_score_gemma":0.004365469,"threshold_uncertainty_score":0.024652421},"labels":[],"label_agreement":null},{"id":"W4280509609","doi":"10.1098/rspb.2021.2721","title":"Past and future uses of text mining in ecology and evolution","year":2022,"lang":"en","type":"review","venue":"Proceedings of the Royal Society B Biological Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":40,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Medical Research Council","keywords":"Computer science; Data science; Evolutionary ecology; Artificial intelligence; Domain (mathematical analysis); Ecology; Biology","score_opus":0.03889944345028952,"score_gpt":0.29404271935548343,"score_spread":0.2551432759051939,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4280509609","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000111527566,0.99520737,0.00049750914,0.003006597,0.00034285648,0.000009273692,0.000022573608,0.000012581862,0.00078967394],"genre_scores_gemma":[0.001017817,0.99496365,0.0016505658,0.0015958755,0.0003990771,0.00002017512,0.00003934577,0.000008460964,0.00030505314],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99815685,0.000753519,0.0003096709,0.00020640179,0.00049458514,0.00007881271],"domain_scores_gemma":[0.9775056,0.018026067,0.0010568856,0.0004274339,0.0025496788,0.00043431367],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.009272782,0.0007494025,0.0015493702,0.006480633,0.0004962847,0.003081355,0.0016231595,0.002502774,0.0025965848],"category_scores_gemma":[0.013541708,0.00044188977,0.0012044774,0.0068781856,0.0024220021,0.0057298117,0.0014766087,0.0030894892,0.0011837053],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004227745,0.000035184537,0.00025107933,0.01955584,0.00012286233,0.00010591547,0.00023379322,0.00028011634,0.00039820533,0.012319109,0.016855778,0.9497998],"study_design_scores_gemma":[0.000024247089,0.00005609231,0.0014228533,0.026031919,0.00019368122,0.00082007097,0.0002853611,0.00023522151,0.00044554178,0.01615021,0.954289,0.00004596122],"about_ca_topic_score_codex":0.002188461,"about_ca_topic_score_gemma":0.0037453724,"teacher_disagreement_score":0.99072725,"about_ca_system_score_codex":0.0018585798,"about_ca_system_score_gemma":0.0048091523,"threshold_uncertainty_score":0.04903978},"labels":[],"label_agreement":null},{"id":"W4281484411","doi":"10.3233/shti220547","title":"An Agile Approach to Accelerate Development and Adoption of Electronic Product Information Standards","year":2022,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Pfizer (Canada)","funders":"","keywords":"Interoperability; Agile software development; Product (mathematics); New product development; Computer science; Knowledge management; Process management; Business; World Wide Web; Software engineering; Marketing","score_opus":0.027842169279354102,"score_gpt":0.3384266786843587,"score_spread":0.3105845094050046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281484411","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020669624,0.00036343082,0.9631521,0.0034995708,0.00017801515,0.0013230798,0.00010693337,0.0017208891,0.008986412],"genre_scores_gemma":[0.055316847,0.00032040238,0.93881387,0.00050937914,0.00004371076,0.0008697931,0.00039544667,0.00049882894,0.0032318614],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.92696226,0.039864227,0.0051743845,0.0048913807,0.020011205,0.0030965153],"domain_scores_gemma":[0.8845402,0.03169228,0.0063013053,0.039418798,0.0331082,0.0049391915],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.080925964,0.0016482421,0.0006452614,0.003880016,0.0027954625,0.009737332,0.0053209094,0.0025627948,0.0022435042],"category_scores_gemma":[0.0848456,0.001715913,0.0019059234,0.0030146025,0.004897582,0.01077269,0.015844624,0.007552626,0.0017804132],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002480135,0.001292982,0.013603694,0.0013478019,0.00031714712,0.0016966485,0.03532768,0.018122887,0.023260018,0.25395384,0.017365936,0.63346344],"study_design_scores_gemma":[0.00034369642,0.0019021146,0.006522548,0.0027715804,0.00032720185,0.0030666078,0.018382695,0.094442554,0.05677677,0.27466473,0.5404006,0.00039878627],"about_ca_topic_score_codex":0.0026688334,"about_ca_topic_score_gemma":0.0027955633,"teacher_disagreement_score":0.080925964,"about_ca_system_score_codex":0.0034395207,"about_ca_system_score_gemma":0.014081788,"threshold_uncertainty_score":0.4279825},"labels":[],"label_agreement":null},{"id":"W4281493375","doi":"10.3233/shti220403","title":"Pretrained Neural Networks Accurately Identify Cancer Recurrence in Medical Record","year":2022,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"CancerCare Manitoba; University of Waterloo","funders":"","keywords":"Medical record; Computer science; Task (project management); Cancer; Artificial intelligence; Disease; Colorectal cancer; Domain (mathematical analysis); Breast cancer; Natural language processing; Medicine; Machine learning; Surgery; Internal medicine","score_opus":0.06847369989992134,"score_gpt":0.41685286299102065,"score_spread":0.3483791630910993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281493375","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81870276,0.004169132,0.15608212,0.0024245572,0.0005450992,0.00017580912,0.005422406,0.004250923,0.0082271155],"genre_scores_gemma":[0.961623,0.0005667364,0.028091377,0.00038906114,0.0001251373,0.000082335224,0.0058993143,0.000058514128,0.0031644977],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995826,0.000073110365,0.00004086043,0.00017823017,0.000059001803,0.00006620409],"domain_scores_gemma":[0.9978561,0.0013831863,0.00021235256,0.00014889918,0.0003468037,0.000052733467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010705283,0.000803255,0.00047758588,0.0007368915,0.00033093328,0.0007613907,0.0007655162,0.00095537095,0.0015294702],"category_scores_gemma":[0.005653032,0.00029910234,0.00053546834,0.00078098045,0.000226462,0.0012941136,0.00048513734,0.0011189918,0.00081974966],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006480187,0.0007717048,0.059460774,0.00025052225,0.00032826178,0.00042458158,0.00018806352,0.45254275,0.0053843916,0.0014072895,0.018336717,0.46025684],"study_design_scores_gemma":[0.000010699411,0.00006241225,0.0046815993,0.000023475903,0.000037013808,0.000037540867,0.000035890236,0.9922334,0.0010416073,0.0010534361,0.000774385,0.000008634143],"about_ca_topic_score_codex":0.020768996,"about_ca_topic_score_gemma":0.02800461,"teacher_disagreement_score":0.020768996,"about_ca_system_score_codex":0.0010406878,"about_ca_system_score_gemma":0.0010389269,"threshold_uncertainty_score":0.041296244},"labels":[],"label_agreement":null},{"id":"W4281565908","doi":"10.7554/elife.70780.sa2","title":"Author response: The LOTUS initiative for open knowledge management in natural products research","year":2022,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"","keywords":"Lotus; Natural (archaeology); Business; Computer science; Geography; Biology; Botany","score_opus":0.2498186064597809,"score_gpt":0.4993427278052674,"score_spread":0.2495241213454865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281565908","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00070610025,0.0004638093,0.0014373144,0.89405733,0.076052554,0.000262508,0.0017937215,0.0019715042,0.023255104],"genre_scores_gemma":[0.006189959,0.0005457056,0.0016790684,0.6803491,0.03106854,0.0005441904,0.0014237652,0.0007481841,0.27745152],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9754453,0.004683796,0.0013453653,0.001698228,0.0144883795,0.0023388881],"domain_scores_gemma":[0.87635475,0.041570313,0.006252145,0.004851888,0.055258054,0.01571281],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.019055855,0.0010414182,0.0017151127,0.0027350294,0.0051215175,0.0097029265,0.0027658618,0.027020369,0.1246248],"category_scores_gemma":[0.10659352,0.0008547926,0.0012761026,0.00173655,0.002862602,0.004702093,0.0051024873,0.017677289,0.07126869],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017746803,0.000006423219,0.00006589788,0.000027597815,0.000002932245,0.00004447183,0.000031616553,0.000011096645,0.000060865543,0.00051865587,0.99763596,0.0015766761],"study_design_scores_gemma":[0.0000474238,0.000021104097,0.0009121004,0.00008714989,0.000008654069,0.00008962652,0.00024700392,0.0002181938,0.00021646463,0.000900277,0.9972052,0.000046674893],"about_ca_topic_score_codex":0.008417187,"about_ca_topic_score_gemma":0.01709993,"teacher_disagreement_score":0.99723417,"about_ca_system_score_codex":0.0052476455,"about_ca_system_score_gemma":0.013335364,"threshold_uncertainty_score":0.41691148},"labels":[],"label_agreement":null},{"id":"W4281566145","doi":"10.3233/shti220603","title":"Designing an Optimal Expansion Method to Improve the Recall of a Genomic Variant Curation-Support Service","year":2022,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; Horizon 2020 Framework Programme","keywords":"Recall; Computer science; Identifier; Precision and recall; Service (business); Information retrieval; Protocol (science); Genome; World Wide Web; Computational biology; Gene; Biology; Genetics; Medicine; Computer network; Psychology","score_opus":0.04244350982157586,"score_gpt":0.3716608041929925,"score_spread":0.32921729437141667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281566145","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17242466,0.0045717005,0.7211534,0.002268331,0.00076966925,0.0014193207,0.004582686,0.08674029,0.0060699224],"genre_scores_gemma":[0.30604103,0.0007366492,0.67706263,0.00071841833,0.000327949,0.0007499343,0.007945609,0.0028335575,0.003584283],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99206245,0.00205772,0.0010139456,0.001521619,0.0024204294,0.00092390465],"domain_scores_gemma":[0.9845798,0.008114575,0.0005072434,0.001954713,0.0044451994,0.00039850804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007793243,0.002522764,0.00263,0.0052250903,0.0012364226,0.0029384578,0.003075241,0.003578692,0.005519876],"category_scores_gemma":[0.037024584,0.0008319573,0.0017258371,0.0033408804,0.0009272547,0.0047752154,0.0023682015,0.0015991697,0.007603333],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025072019,0.00078486506,0.021322636,0.0015127545,0.00033149618,0.0006502697,0.00072242157,0.027534999,0.06991956,0.0032624824,0.051594544,0.8198568],"study_design_scores_gemma":[0.0003824089,0.0004423497,0.006795438,0.00017176404,0.00035446885,0.0011721151,0.00062846625,0.8941069,0.06888121,0.0058080796,0.021120882,0.00013582835],"about_ca_topic_score_codex":0.009424263,"about_ca_topic_score_gemma":0.007006247,"teacher_disagreement_score":0.009424263,"about_ca_system_score_codex":0.001225194,"about_ca_system_score_gemma":0.003530311,"threshold_uncertainty_score":0.04121512},"labels":[],"label_agreement":null},{"id":"W4281856165","doi":"10.1101/2022.05.24.22275490","title":"Evaluating the impact on clinical task efficiency of a natural language processing algorithm for searching medical documents: Prospective crossover study","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Global Health Research","funders":"University of Edinburgh; UK Research and Innovation","keywords":"Computer science; Task (project management); Workflow; Upload; Artificial intelligence; Natural language processing; Information retrieval; String searching algorithm; Search algorithm; Machine learning; Algorithm; World Wide Web; Pattern matching; Database","score_opus":0.05282356704751386,"score_gpt":0.4958383169671636,"score_spread":0.44301474991964973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4281856165","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99920195,0.000059633603,0.00023305278,0.000009547001,0.000013750527,0.00033575157,0.000034461806,0.0000059983386,0.00010586394],"genre_scores_gemma":[0.9973488,0.000055223456,0.0011320679,0.000035415418,0.00002161464,0.0009050123,0.00010366246,0.000007733797,0.00039050123],"study_design_codex":"randomized_trial","study_design_gemma":"nonrandomized_trial","domain_scores_codex":[0.9940276,0.0034116441,0.0006326653,0.0008837611,0.00062485103,0.00041949694],"domain_scores_gemma":[0.9795873,0.0112511385,0.0037337244,0.0017906459,0.00171255,0.0019246887],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011093133,0.0008298054,0.0011888909,0.0005603026,0.00066561415,0.0010780385,0.0006478198,0.0010278608,0.002589974],"category_scores_gemma":[0.017962156,0.0005971071,0.0012153167,0.00045420008,0.0010529491,0.0010683979,0.0005947741,0.001125459,0.0005182815],"study_design_candidate":"nonrandomized_trial","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.4123955,0.23511499,0.2701308,0.0010958031,0.0022251979,0.00063025294,0.00682501,0.002008257,0.016603692,0.00033942642,0.0011151212,0.05151594],"study_design_scores_gemma":[0.016321395,0.7452134,0.23104689,0.00005272187,0.0005523805,0.00014053709,0.0009068479,0.0017423106,0.0030280487,0.00015080492,0.0007653322,0.00007920258],"about_ca_topic_score_codex":0.00064488116,"about_ca_topic_score_gemma":0.00057392695,"teacher_disagreement_score":0.98890686,"about_ca_system_score_codex":0.00061261974,"about_ca_system_score_gemma":0.0007990169,"threshold_uncertainty_score":0.058666825},"labels":[],"label_agreement":null},{"id":"W4283216500","doi":"10.3389/fnut.2022.928837","title":"Establishing a Common Nutritional Vocabulary - From Food Production to Diet","year":2022,"lang":"en","type":"article","venue":"Frontiers in Nutrition","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":18,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development; Biotechnology and Biological Sciences Research Council; Southern Cross University; U.S. Department of Agriculture","keywords":"Standardization; Production (economics); Vocabulary; Food security; Food processing; Ontology; Consumption (sociology); Computer science; Business; Food science; Agriculture; Biology; Ecology; Sociology; Economics","score_opus":0.011080378933345441,"score_gpt":0.23364117314696772,"score_spread":0.2225607942136223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283216500","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019599963,0.0017297409,0.9235233,0.0070058447,0.0007211572,0.0008724654,0.01740706,0.0021769106,0.026963504],"genre_scores_gemma":[0.11883589,0.0028898693,0.8433043,0.002143003,0.00027421152,0.0014620102,0.025530253,0.00077472295,0.0047858036],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963438,0.0008847696,0.0009909238,0.0007995881,0.00083777454,0.0001431078],"domain_scores_gemma":[0.9936841,0.002248174,0.0007652273,0.0015463658,0.0014430116,0.00031300186],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006710703,0.00087871164,0.0010298118,0.006126319,0.0017342989,0.0056480994,0.002340679,0.0016160505,0.0025344386],"category_scores_gemma":[0.0123387445,0.00062271906,0.0016926741,0.005513453,0.003979018,0.0097518675,0.0038227565,0.0032213696,0.0012283869],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013146276,0.00012039587,0.0061994586,0.0014638279,0.00013788023,0.0003606818,0.0042066877,0.0034402553,0.011106487,0.8394022,0.018968448,0.11446232],"study_design_scores_gemma":[0.000043590597,0.00006935722,0.0044951374,0.001478799,0.00016788684,0.00062963186,0.0032603282,0.011081355,0.0057553393,0.42574853,0.54717016,0.00009988479],"about_ca_topic_score_codex":0.011913989,"about_ca_topic_score_gemma":0.011995584,"teacher_disagreement_score":0.011913989,"about_ca_system_score_codex":0.0026824113,"about_ca_system_score_gemma":0.008304246,"threshold_uncertainty_score":0.035490036},"labels":[],"label_agreement":null},{"id":"W4283327249","doi":"10.3390/fi14070190","title":"First Steps of Asthma Management with a Personalized Ontology Model","year":2022,"lang":"en","type":"article","venue":"Future Internet","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Chicoutimi","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Université du Québec à Chicoutimi","keywords":"Asthma; Ontology; Computer science; Wheeze; Asthma management; Semantic Web; Medicine; Intensive care medicine; World Wide Web; Immunology","score_opus":0.008311099810004668,"score_gpt":0.22676183646879658,"score_spread":0.21845073665879192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283327249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012764711,0.00067410234,0.959803,0.0029970836,0.0001366805,0.00063720584,0.0010762041,0.002748274,0.019162735],"genre_scores_gemma":[0.107513264,0.000870847,0.88294744,0.00047232353,0.000049038128,0.00034954594,0.0020486366,0.00015889369,0.0055900314],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989171,0.00027698278,0.00011124589,0.00018901017,0.00040500934,0.000100681296],"domain_scores_gemma":[0.99929976,0.00020992983,0.000058408932,0.00019940575,0.00017980309,0.000052697822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013530144,0.0005112192,0.00050558976,0.0010106792,0.0008790553,0.0027002823,0.0014505998,0.0010639395,0.0041177515],"category_scores_gemma":[0.0024944867,0.0004259822,0.0016557068,0.00085967313,0.0004552643,0.002652131,0.0019905213,0.0015156578,0.001480372],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024144712,0.00079245644,0.0070598884,0.00065937,0.0003386527,0.0019450831,0.0015261957,0.0652356,0.012836424,0.30742982,0.022983886,0.5789512],"study_design_scores_gemma":[0.000068632035,0.00012918565,0.0023616913,0.0003484372,0.00034936456,0.00099633,0.0006962755,0.56216586,0.009957725,0.24716884,0.17566343,0.00009427388],"about_ca_topic_score_codex":0.008434421,"about_ca_topic_score_gemma":0.01021399,"teacher_disagreement_score":0.008434421,"about_ca_system_score_codex":0.0009988378,"about_ca_system_score_gemma":0.0032118126,"threshold_uncertainty_score":0.01677066},"labels":[],"label_agreement":null},{"id":"W4283364896","doi":"10.1101/2022.02.16.22268694","title":"Automated identification of unstandardized medication data: A scalable and flexible data standardization pipeline using RxNorm on GEMINI multicenter hospital data","year":2022,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto; St. Michael's Hospital","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research; Alliance de recherche numérique du Canada; Canadian Frailty Network; University of Toronto; University Health Network","keywords":"Standardization; Computer science; Identifier; Pharmacy; False positive paradox; Data mining; Coding (social sciences); Information retrieval; Medicine; Artificial intelligence; Statistics; Mathematics; Family medicine","score_opus":0.07059790562919259,"score_gpt":0.36991126172167255,"score_spread":0.29931335609247994,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283364896","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.119767755,0.0016045481,0.54291344,0.00336185,0.00020428849,0.0022392392,0.044307202,0.27928486,0.006316784],"genre_scores_gemma":[0.21796466,0.00051392894,0.7088275,0.00080245594,0.00010561164,0.0007956206,0.06495728,0.004092445,0.0019404938],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9845441,0.0038055233,0.002504772,0.00397389,0.0046987236,0.00047294697],"domain_scores_gemma":[0.9645962,0.011768664,0.004757347,0.010093509,0.007982344,0.00080203096],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022409724,0.0016569289,0.001209887,0.0068685347,0.0011764152,0.004009985,0.0023499976,0.0007028171,0.0019007156],"category_scores_gemma":[0.045012947,0.00094491703,0.0016153859,0.005283065,0.0010861969,0.004568312,0.004704499,0.0015650525,0.001714444],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019140178,0.0006043417,0.13493775,0.0015309454,0.00087829365,0.0011471376,0.0041915993,0.021369655,0.04694799,0.010224589,0.10238338,0.67387027],"study_design_scores_gemma":[0.00060146925,0.0006943723,0.099302255,0.00066953903,0.00040768727,0.0012305431,0.0027097252,0.52554005,0.16602717,0.021280322,0.18096231,0.000574565],"about_ca_topic_score_codex":0.020570207,"about_ca_topic_score_gemma":0.01604202,"teacher_disagreement_score":0.022409724,"about_ca_system_score_codex":0.002587423,"about_ca_system_score_gemma":0.006839572,"threshold_uncertainty_score":0.11851537},"labels":[],"label_agreement":null},{"id":"W4283523834","doi":"10.1016/j.jval.2022.04.1092","title":"HSD92 Natural Language Processing (NLP) of Unstructured EMR Data to Describe Treatment Patterns and Achievement of Guideline Recommended LDL-C Targets of Canadian Patients with Ascvd and Hypercholesterolemia","year":2022,"lang":"en","type":"article","venue":"Value in Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Novartis (Canada); University of British Columbia; University of Regina; St. Michael's Hospital","funders":"","keywords":"Ezetimibe; Statin; Medicine; Guideline; Internal medicine; Medical record; Electronic medical record; Family medicine; Pathology","score_opus":0.03963855517563155,"score_gpt":0.29906989038568843,"score_spread":0.2594313352100569,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283523834","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.19921006,0.00066575664,0.0468448,0.0013045901,0.000101263446,0.0012991057,0.72469467,0.013449303,0.012430467],"genre_scores_gemma":[0.30077398,0.0005739716,0.12345828,0.0003854299,0.000038420487,0.0006772513,0.5679687,0.00054123765,0.005582722],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99934727,0.00009428705,0.00010712312,0.00018738776,0.000188822,0.000075087],"domain_scores_gemma":[0.99681866,0.0018122676,0.00021320714,0.0002482448,0.0008157973,0.000091876245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009540107,0.0006313121,0.00033104562,0.002691046,0.00062853174,0.0010665425,0.00054893503,0.0003706973,0.0047745514],"category_scores_gemma":[0.006566492,0.0002288213,0.0008180907,0.0025937003,0.0003528027,0.00039040955,0.0005377516,0.0005444494,0.0014677137],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028103953,0.00056077936,0.20175739,0.0046847807,0.00073078234,0.0034840857,0.004735659,0.04127975,0.041422877,0.010945512,0.22981222,0.4577758],"study_design_scores_gemma":[0.0005384641,0.0003855801,0.33849552,0.00054591504,0.00069589925,0.0019021216,0.003989997,0.20644219,0.050638244,0.01072043,0.38538402,0.00026157667],"about_ca_topic_score_codex":0.46628997,"about_ca_topic_score_gemma":0.52558374,"teacher_disagreement_score":0.53371,"about_ca_system_score_codex":0.0031721096,"about_ca_system_score_gemma":0.008978527,"threshold_uncertainty_score":0.92715174},"labels":[],"label_agreement":null},{"id":"W4283589003","doi":"10.2196/37817","title":"A Syntactic Information–Based Classification Model for Medical Literature: Algorithm Development and Validation Study","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Fundamental Research Funds for the Central Universities; Natural Science Foundation of Liaoning Province; National Natural Science Foundation of China","keywords":"Computer science; Algorithm; Data mining; Information model; Artificial intelligence; Natural language processing; Software engineering","score_opus":0.024484167493020438,"score_gpt":0.31115051478538613,"score_spread":0.2866663472923657,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4283589003","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09974271,0.0026259725,0.8834963,0.0019943824,0.00022393299,0.0010672839,0.001459749,0.006167162,0.0032224972],"genre_scores_gemma":[0.3702065,0.0011752899,0.61817104,0.0006391226,0.00017941835,0.0016009151,0.004587725,0.00020541009,0.003234559],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984958,0.00046127324,0.00020235172,0.00035860637,0.00036715,0.000114762945],"domain_scores_gemma":[0.99357325,0.00385669,0.0002622899,0.0003282857,0.0018737,0.000105859246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036671027,0.0013331851,0.0013157598,0.0039086924,0.0008407592,0.0017049919,0.0023305325,0.0023799513,0.003380347],"category_scores_gemma":[0.011254696,0.0004288641,0.0015040654,0.002814554,0.0005075291,0.0021470995,0.0012875115,0.0019128815,0.0016250149],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040153708,0.00065716624,0.01234859,0.0003257049,0.00037196383,0.00027738736,0.00014033854,0.1991049,0.004181137,0.0035965212,0.010806183,0.7677886],"study_design_scores_gemma":[0.00002330636,0.000033967972,0.0004418455,0.000019515823,0.00003842177,0.000053611602,0.00002197588,0.99686074,0.00075394486,0.0011965345,0.0005497524,0.0000064292553],"about_ca_topic_score_codex":0.015604333,"about_ca_topic_score_gemma":0.01298924,"teacher_disagreement_score":0.015604333,"about_ca_system_score_codex":0.002269108,"about_ca_system_score_gemma":0.0039724847,"threshold_uncertainty_score":0.031027019},"labels":[],"label_agreement":null},{"id":"W4285236424","doi":"10.1007/978-3-031-09342-5_3","title":"A Knowledge Graph Completion Method Applied to Literature-Based Discovery for Predicting Missing Links Targeting Cancer Drug Repurposing","year":2022,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Repurposing; Computer science; Drug repositioning; Knowledge graph; Knowledge extraction; Graph; Drug discovery; Artificial intelligence; Information retrieval; Theoretical computer science; Drug; Bioinformatics; Medicine","score_opus":0.01953491837939493,"score_gpt":0.30579614245022957,"score_spread":0.28626122407083465,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285236424","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02202531,0.0006031698,0.9670738,0.00026613873,0.00011095,0.00031303943,0.0029338943,0.004430143,0.0022435954],"genre_scores_gemma":[0.09572325,0.00036659857,0.89365935,0.000087577515,0.00006515288,0.000219423,0.005828446,0.00024650214,0.003803677],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99933356,0.00010776679,0.000058610727,0.00020233385,0.00025392562,0.000043659722],"domain_scores_gemma":[0.99786276,0.0011781114,0.00012668004,0.00023964887,0.00049239065,0.0001004394],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010222767,0.0007427931,0.0010130507,0.0051582893,0.0008381882,0.001110332,0.0014800867,0.0010333728,0.0047492976],"category_scores_gemma":[0.0039865826,0.00037402945,0.0014599584,0.004132357,0.0004851362,0.0012894142,0.0012636026,0.0010719019,0.0016905842],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029624568,0.0004680433,0.0034036736,0.00039494698,0.00026232476,0.00043265364,0.00017749426,0.08370095,0.010188109,0.009356873,0.018928863,0.8723899],"study_design_scores_gemma":[0.00003825966,0.00007029851,0.0012136768,0.000037536938,0.00007632273,0.00017809465,0.000070649075,0.9772092,0.002996064,0.012373778,0.0057105385,0.000025640042],"about_ca_topic_score_codex":0.016483806,"about_ca_topic_score_gemma":0.018151797,"teacher_disagreement_score":0.016483806,"about_ca_system_score_codex":0.0005009467,"about_ca_system_score_gemma":0.0022712217,"threshold_uncertainty_score":0.03277576},"labels":[],"label_agreement":null},{"id":"W4285605376","doi":"10.1016/j.jbi.2022.104134","title":"Call for papers: Semantics-enabled biomedical literature analytics","year":2022,"lang":"en","type":"paratext","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Analytics; Data science; Semantics (computer science); Health informatics; Information retrieval; World Wide Web; Programming language; Medicine","score_opus":0.01480760056961364,"score_gpt":0.2913976498659923,"score_spread":0.27659004929637865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4285605376","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008878569,0.007206669,0.4532101,0.019334277,0.0070733815,0.0011229176,0.19576754,0.23678267,0.070623845],"genre_scores_gemma":[0.057356253,0.006750494,0.40959227,0.0041508367,0.0030674436,0.0013262266,0.38981533,0.02915739,0.09878377],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979936,0.00041307803,0.00023202358,0.0003002325,0.00097532914,0.00008583523],"domain_scores_gemma":[0.98920923,0.005077966,0.0004641769,0.00253443,0.0012608675,0.0014533631],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.003293814,0.00239626,0.0012390693,0.0104212,0.0013829996,0.008724269,0.0023334746,0.002791877,0.1098095],"category_scores_gemma":[0.02284178,0.0009245374,0.0016900083,0.009761169,0.0008431387,0.0100633735,0.0079442635,0.002120098,0.06567553],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032692205,0.00020453615,0.0010428185,0.001446953,0.00021326604,0.0005011015,0.0002951681,0.0021982265,0.004132808,0.048256703,0.66278166,0.27859974],"study_design_scores_gemma":[0.0002452565,0.00008136711,0.0010578316,0.00029374234,0.0001077344,0.0004325673,0.00022806437,0.02657127,0.0052478323,0.11173885,0.8539148,0.00008070396],"about_ca_topic_score_codex":0.0016955683,"about_ca_topic_score_gemma":0.0021819398,"teacher_disagreement_score":0.8901905,"about_ca_system_score_codex":0.0008656704,"about_ca_system_score_gemma":0.0018774056,"threshold_uncertainty_score":0.3673494},"labels":[],"label_agreement":null},{"id":"W4286974625","doi":"","title":"An Ontological Analysis of Health Procedure Information","year":2021,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Computer science; Information retrieval; Data science; Natural language processing","score_opus":0.015844789934634076,"score_gpt":0.2714375917902347,"score_spread":0.2555928018556006,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286974625","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.052954827,0.0010839514,0.90638095,0.00495038,0.00024027778,0.00070436526,0.010405363,0.002279207,0.021000806],"genre_scores_gemma":[0.32654303,0.0013393095,0.648708,0.0007433888,0.00016786932,0.00034293547,0.015900778,0.00042309496,0.0058315457],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9954651,0.0009855207,0.0005892866,0.0006868904,0.0019071357,0.00036611172],"domain_scores_gemma":[0.9925654,0.0034186707,0.00057753565,0.0013586963,0.0017790081,0.0003007323],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039863004,0.0005932005,0.00066182675,0.010100618,0.0024350155,0.005939729,0.001481156,0.0011166848,0.004131852],"category_scores_gemma":[0.012611718,0.00049701944,0.0032264981,0.008462407,0.0016903115,0.0069894404,0.0030787443,0.0016984312,0.00073408266],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033461527,0.0003553831,0.014073079,0.000998319,0.00028942106,0.001488612,0.0043758918,0.014896825,0.009813186,0.66489804,0.01526571,0.27321097],"study_design_scores_gemma":[0.000047135945,0.00009511149,0.012290795,0.0008963391,0.0007393557,0.0011635487,0.0047289464,0.2187296,0.012567502,0.5827884,0.16584189,0.00011142069],"about_ca_topic_score_codex":0.018976714,"about_ca_topic_score_gemma":0.020475052,"teacher_disagreement_score":0.018976714,"about_ca_system_score_codex":0.0033403276,"about_ca_system_score_gemma":0.005050041,"threshold_uncertainty_score":0.03773254},"labels":[],"label_agreement":null},{"id":"W4287118057","doi":"10.23889/ijpds.v5i1.1362","title":"Concept libraries for automatic electronic health record based phenotyping: A review.","year":2021,"lang":"en","type":"review","venue":"PEARL (University of Plymouth)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Resource (disambiguation); Data science; Disease; Health records; Process (computing); Reuse; Electronic health record; Medicine; Health care; Political science; Engineering","score_opus":0.037447964569957116,"score_gpt":0.3057809576412142,"score_spread":0.2683329930712571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287118057","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00058494357,0.99135476,0.0037129754,0.001353184,0.00024335686,0.0005206428,0.0004305671,0.000100203164,0.00169936],"genre_scores_gemma":[0.0077792206,0.96756804,0.020394063,0.0013953182,0.00015998262,0.0012882635,0.0009289289,0.00003115634,0.00045503347],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.9909622,0.0037033926,0.0023336029,0.0005645615,0.0022224563,0.00021385803],"domain_scores_gemma":[0.9138085,0.07031232,0.0073462534,0.0013293386,0.006595333,0.0006081969],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.019496446,0.0013268971,0.0030242975,0.026676826,0.00095465616,0.004618347,0.0034628573,0.0018199232,0.0061206543],"category_scores_gemma":[0.05675912,0.0010504778,0.0039531197,0.020375954,0.0019693503,0.008662937,0.0037335681,0.0021705974,0.0014361187],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013413152,0.00009161904,0.0006480788,0.22220439,0.0006205557,0.00012072579,0.0006753023,0.00044665186,0.00018420655,0.008176831,0.013555378,0.75314206],"study_design_scores_gemma":[0.00015059605,0.00031916233,0.0036629483,0.55476815,0.0036268432,0.0015063909,0.0010614686,0.0009178425,0.00079179165,0.009409917,0.42360514,0.0001797457],"about_ca_topic_score_codex":0.0050323997,"about_ca_topic_score_gemma":0.008326458,"teacher_disagreement_score":0.026676826,"about_ca_system_score_codex":0.004516537,"about_ca_system_score_gemma":0.015312471,"threshold_uncertainty_score":0.10310829},"labels":[],"label_agreement":null},{"id":"W4287601596","doi":"10.5281/zenodo.4482922","title":"MeDAL: Medical Abbreviation Disambiguation Dataset for Natural Language Understanding Pretraining","year":2020,"lang":"en","type":"paratext","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Natural language processing; Medal; Artificial intelligence; Natural language; Natural (archaeology); Information retrieval; History","score_opus":0.07407423310831963,"score_gpt":0.3220944828030508,"score_spread":0.2480202496947312,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4287601596","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009205873,0.0013088614,0.0062352396,0.0015052364,0.00046674497,0.0007342814,0.96557873,0.008515803,0.0064491974],"genre_scores_gemma":[0.006547213,0.00016659073,0.010000189,0.0005370414,0.000050767685,0.0006177048,0.9800149,0.00025699622,0.0018085919],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975963,0.0005512353,0.0004334199,0.0007936361,0.00044935665,0.00017586944],"domain_scores_gemma":[0.99693346,0.0010795224,0.00022901174,0.0006465875,0.0007285353,0.00038294398],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002098877,0.0023568526,0.0012190691,0.0040993392,0.0014903743,0.0013438668,0.0030184903,0.0029825794,0.02802622],"category_scores_gemma":[0.009784666,0.00052431016,0.0016058291,0.0023972017,0.00080977904,0.0017716492,0.0026457945,0.0026198905,0.038036942],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041054623,0.00014790703,0.0022641586,0.0014288996,0.00008148145,0.00034416615,0.00017172692,0.0009887217,0.0023761739,0.000997667,0.9595207,0.031267934],"study_design_scores_gemma":[0.0010723949,0.00036019148,0.018380538,0.00059435196,0.00017010153,0.002118808,0.000631833,0.012619841,0.009988935,0.0046029515,0.9492639,0.00019624026],"about_ca_topic_score_codex":0.013416283,"about_ca_topic_score_gemma":0.029274162,"teacher_disagreement_score":0.02802622,"about_ca_system_score_codex":0.0019459457,"about_ca_system_score_gemma":0.003625447,"threshold_uncertainty_score":0.09375703},"labels":[],"label_agreement":null},{"id":"W4288033978","doi":"10.1002/cjce.24574","title":"Is it time to change how we write scientific articles?","year":2022,"lang":"en","type":"article","venue":"The Canadian Journal of Chemical Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"CLARITY; Perspective (graphical); Focus (optics); Style (visual arts); Norm (philosophy); Computer science; Engineering ethics; Epistemology; Philosophy; Literature; Artificial intelligence; Art; Engineering; Chemistry; Physics","score_opus":0.02637485751147193,"score_gpt":0.22174986206805447,"score_spread":0.19537500455658255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288033978","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013853996,0.006006967,0.024393162,0.9142349,0.04250117,0.00006510975,0.000059612597,0.00032657178,0.011027102],"genre_scores_gemma":[0.12591365,0.01990989,0.0957406,0.6582242,0.06996949,0.00083006296,0.00029520888,0.0021002984,0.027016602],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.85767055,0.08867873,0.01356256,0.0071350634,0.029504178,0.0034489227],"domain_scores_gemma":[0.51931447,0.25771198,0.03321315,0.028731203,0.13391112,0.027118046],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.15393922,0.0008290546,0.0013485824,0.0034990397,0.0110273585,0.043412052,0.004991379,0.009661457,0.0065851817],"category_scores_gemma":[0.33687916,0.0008848985,0.0011668825,0.0032047308,0.036311716,0.033700783,0.0084898705,0.032083385,0.00833672],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016873323,0.00018430394,0.0015556097,0.0017488239,0.00018277274,0.00038654357,0.041201815,0.0005466681,0.0034102267,0.30797648,0.4648372,0.17780083],"study_design_scores_gemma":[0.000048778067,0.000085467866,0.00078890583,0.0017376025,0.00006572971,0.00034193665,0.017397413,0.00032224794,0.0009041926,0.14224231,0.8359159,0.00014944235],"about_ca_topic_score_codex":0.0055987556,"about_ca_topic_score_gemma":0.0054585575,"teacher_disagreement_score":0.84606075,"about_ca_system_score_codex":0.00996833,"about_ca_system_score_gemma":0.024649352,"threshold_uncertainty_score":0.81411815},"labels":[],"label_agreement":null},{"id":"W4288257343","doi":"10.5281/zenodo.3374904","title":"MatthewBCooke/Pathfinder: First public release","year":2019,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Pathfinder; Political science; Computer science; Library science","score_opus":0.027113680021617265,"score_gpt":0.2343335125461407,"score_spread":0.20721983252452345,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4288257343","genre_codex":"dataset","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017389109,0.0004642613,0.012652305,0.000543815,0.00032019886,0.00016949217,0.7435553,0.22756127,0.012994505],"genre_scores_gemma":[0.009863928,0.0006075715,0.026606066,0.0003450516,0.000116294934,0.00053418434,0.85030323,0.09433956,0.01728413],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989973,0.000065876775,0.000091563736,0.00025294026,0.00042350916,0.00016887693],"domain_scores_gemma":[0.99702364,0.0009781398,0.00020524231,0.00085683755,0.00061716605,0.0003189097],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015292751,0.0023841804,0.0016816797,0.005432544,0.00069008034,0.002664993,0.003239641,0.0018571536,0.34563634],"category_scores_gemma":[0.009663789,0.001712533,0.0019337244,0.00582042,0.0005430911,0.003240223,0.0030080988,0.0016979737,0.2502251],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033222008,0.00005584608,0.0007786413,0.0010371166,0.000078444595,0.00009726864,0.00009998238,0.000299937,0.0017020744,0.00096103305,0.96277386,0.031783678],"study_design_scores_gemma":[0.000553636,0.00007039759,0.0047609396,0.0003951456,0.000086760636,0.00029760852,0.000090626905,0.0032009892,0.0046997834,0.00493847,0.9807858,0.000119855904],"about_ca_topic_score_codex":0.009583887,"about_ca_topic_score_gemma":0.01376907,"teacher_disagreement_score":0.34563634,"about_ca_system_score_codex":0.0008664787,"about_ca_system_score_gemma":0.0023316604,"threshold_uncertainty_score":0.93337035},"labels":[],"label_agreement":null},{"id":"W4289667473","doi":"10.1101/2022.08.02.502449","title":"Using language models and ontology topology to perform semantic mapping of traits between biomedical datasets","year":2022,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Future Earth","funders":"Medical Research Council; University of Bristol","keywords":"Computer science; Biobank; Ontology; Pairwise comparison; Matching (statistics); Natural language processing; Phenome; Trait; Information retrieval; Semantic similarity; Artificial intelligence; Code (set theory); Bioinformatics; Set (abstract data type); Biology","score_opus":0.03905520679226976,"score_gpt":0.28979941301322715,"score_spread":0.2507442062209574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4289667473","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1969794,0.0018334065,0.7575288,0.0038056297,0.00041196705,0.00042091165,0.012839737,0.020807106,0.00537304],"genre_scores_gemma":[0.59737116,0.0005844627,0.37472174,0.00066112005,0.00012388574,0.0006110764,0.022805197,0.0013453835,0.0017759603],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99640816,0.0018415812,0.00028084413,0.0009087764,0.00038711054,0.0001735019],"domain_scores_gemma":[0.9831092,0.013392071,0.00091686053,0.0012020268,0.0010449444,0.00033495555],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076669008,0.0014216278,0.00072963577,0.005046733,0.0011566883,0.0032548234,0.001614378,0.0014029505,0.0025634908],"category_scores_gemma":[0.0316916,0.00054715894,0.0027534622,0.002851943,0.0007948069,0.0046402183,0.002755145,0.0021568276,0.0014811292],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013499794,0.00077308953,0.09132033,0.0016464092,0.0016830643,0.0012073645,0.0032504112,0.4306829,0.010956348,0.036449853,0.02750065,0.39317966],"study_design_scores_gemma":[0.000045939825,0.00007486583,0.0034539716,0.000098866825,0.00008841142,0.00015065057,0.00042588287,0.9442293,0.002218728,0.04416337,0.005002931,0.000047093345],"about_ca_topic_score_codex":0.017178785,"about_ca_topic_score_gemma":0.023684023,"teacher_disagreement_score":0.017178785,"about_ca_system_score_codex":0.0023577823,"about_ca_system_score_gemma":0.0025921725,"threshold_uncertainty_score":0.040546954},"labels":[],"label_agreement":null},{"id":"W4290996648","doi":"10.1109/icc45855.2022.9838418","title":"Problem Oriented Medical Translational Services based on GraphQL Connectivity","year":2022,"lang":"en","type":"article","venue":"ICC 2022 - IEEE International Conference on Communications","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Translational research; Computer science; Data science; Bridging (networking); Translational medicine; Health care; World Wide Web; Medical research; Big data; Variety (cybernetics); Medicine; Artificial intelligence; Computer security; Political science","score_opus":0.054449646511853765,"score_gpt":0.3463955132168181,"score_spread":0.29194586670496436,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4290996648","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013072526,0.0010753742,0.82455033,0.0060767597,0.0005469827,0.0009843116,0.008225396,0.123067506,0.022400776],"genre_scores_gemma":[0.38618645,0.002143421,0.54183203,0.0059713367,0.00051482767,0.0015692974,0.038322322,0.008963262,0.01449705],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9947107,0.0016261045,0.00081667746,0.00090237445,0.0015003869,0.0004438363],"domain_scores_gemma":[0.9936767,0.0021831696,0.00035820968,0.002391066,0.0008463586,0.00054447807],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004964046,0.00097563484,0.0009513634,0.0019946233,0.0012045227,0.005878145,0.0037485622,0.0020576564,0.011834588],"category_scores_gemma":[0.013414076,0.0005699687,0.002010897,0.0024873693,0.0013387015,0.0072100917,0.008857531,0.0028339212,0.0041078064],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001820922,0.0008397792,0.0063877474,0.0026795855,0.00052512815,0.002652113,0.0031265628,0.038240924,0.018431418,0.36663738,0.19049494,0.3681635],"study_design_scores_gemma":[0.00038968126,0.00023820976,0.0017158694,0.00029968403,0.00016057934,0.0012127982,0.0009452554,0.29145518,0.015405649,0.31888562,0.36908045,0.00021109023],"about_ca_topic_score_codex":0.006740739,"about_ca_topic_score_gemma":0.003922546,"teacher_disagreement_score":0.011834588,"about_ca_system_score_codex":0.0017706242,"about_ca_system_score_gemma":0.0030903309,"threshold_uncertainty_score":0.039590657},"labels":[],"label_agreement":null},{"id":"W4292587986","doi":"","title":"Combiner plongements de graphes et clustering pour l'alignement de connaissances pharmacogénomiques","year":2022,"lang":"fr","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ubisoft (Canada)","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Korea Institute of Public Finance; University of Cambridge; Institute for Catastrophic Loss Reduction","keywords":"Computer science; Cluster analysis; Humanities; Artificial intelligence; Art","score_opus":0.025158852811817702,"score_gpt":0.28729133217588165,"score_spread":0.262132479364064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4292587986","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07819674,0.0020295973,0.8939953,0.00092586095,0.00029101112,0.00034992647,0.009019457,0.0125091225,0.0026829704],"genre_scores_gemma":[0.25868955,0.0012569416,0.7076588,0.00030314963,0.00019394055,0.00037988432,0.023226174,0.0018202999,0.006471258],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971698,0.0006485325,0.00022310173,0.00093566003,0.0008314396,0.00019146632],"domain_scores_gemma":[0.99396294,0.0032790888,0.0004775983,0.0009047014,0.0010992982,0.0002762456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020758011,0.0021214737,0.0015046549,0.011301535,0.0010425976,0.0034103603,0.0013072566,0.0018847899,0.0041546714],"category_scores_gemma":[0.0090963645,0.0009276965,0.0035250778,0.008033241,0.000719354,0.0027275237,0.0016885936,0.0016485999,0.0026047768],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014846899,0.0006765871,0.031255286,0.0013651965,0.002385816,0.00093500805,0.0009873996,0.12719488,0.041402083,0.017094122,0.0210995,0.7541195],"study_design_scores_gemma":[0.00016838248,0.0003344211,0.02614249,0.00021008764,0.0007726368,0.0007679273,0.0008495581,0.8664242,0.021178922,0.051005326,0.032028332,0.00011778928],"about_ca_topic_score_codex":0.02004097,"about_ca_topic_score_gemma":0.034547787,"teacher_disagreement_score":0.02004097,"about_ca_system_score_codex":0.0012500642,"about_ca_system_score_gemma":0.0015338418,"threshold_uncertainty_score":0.039848626},"labels":[],"label_agreement":null},{"id":"W4293167098","doi":"10.26650/b/et07.2021.003.02","title":"Medikal Ontolojiler","year":2021,"lang":"tr","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mitel (Canada)","funders":"","keywords":"Computer science","score_opus":0.027475173644057046,"score_gpt":0.2662000296090326,"score_spread":0.23872485596497556,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4293167098","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03707958,0.012664398,0.19800724,0.0116118435,0.0069256867,0.00036528063,0.00573212,0.012063561,0.71555024],"genre_scores_gemma":[0.35557488,0.03493806,0.2301352,0.0062509296,0.003405789,0.0008502093,0.015661856,0.0065583787,0.3466247],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9979563,0.00042196296,0.00026903706,0.000581978,0.0005669459,0.0002037625],"domain_scores_gemma":[0.999183,0.0001896094,0.00009543183,0.00018550719,0.00028716604,0.000059270067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015777999,0.0011816589,0.0008544058,0.0020906425,0.0021881608,0.007671212,0.0009302611,0.0012926906,0.04366473],"category_scores_gemma":[0.002335426,0.00084936345,0.0010242884,0.0018387362,0.0029912218,0.008734676,0.0046960493,0.0050546792,0.035265002],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044247735,0.00020070041,0.0012450641,0.0012165771,0.000106542495,0.0018030991,0.0040536816,0.0024178023,0.008799826,0.5363265,0.064993195,0.37839454],"study_design_scores_gemma":[0.00005592505,0.0000663758,0.00091111404,0.000368541,0.000045159974,0.001086542,0.00079459004,0.0017195749,0.0038686707,0.08959201,0.9014377,0.000053805226],"about_ca_topic_score_codex":0.0034794786,"about_ca_topic_score_gemma":0.002434227,"teacher_disagreement_score":0.04366473,"about_ca_system_score_codex":0.0031596348,"about_ca_system_score_gemma":0.002747481,"threshold_uncertainty_score":0.14607304},"labels":[],"label_agreement":null},{"id":"W4294585644","doi":"10.7287/peerj.13843v0.2/reviews/1","title":"Peer Review #1 of \"Non-synonymous to synonymous substitutions suggest that orthologs tend to keep their functions, while paralogs are a source of functional novelty (v0.2)\"","year":2022,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wilfrid Laurier University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Novelty; Evolutionary biology; Biology; Computational biology; Genetics; Psychology; Social psychology","score_opus":0.05890213301775135,"score_gpt":0.30682233504601564,"score_spread":0.24792020202826429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294585644","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0057668434,0.014320139,0.014008128,0.13520397,0.5906938,0.006684761,0.035573337,0.0066550327,0.19109407],"genre_scores_gemma":[0.037983086,0.02176747,0.015325581,0.018899996,0.119859405,0.0034377738,0.05386961,0.007209375,0.72164786],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98450595,0.0026346704,0.0017851953,0.0015024381,0.008102858,0.0014689245],"domain_scores_gemma":[0.8370967,0.0092677325,0.004086961,0.008157565,0.1287564,0.01263455],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.010974007,0.0014725162,0.0028884863,0.0056144753,0.004507289,0.011157868,0.0040382384,0.0038125666,0.39613342],"category_scores_gemma":[0.09942279,0.0012517357,0.002309974,0.003947791,0.0016436662,0.006486087,0.0063838097,0.0028557817,0.31321684],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000043742624,0.000014009613,0.00042731306,0.000953919,0.000025956622,0.000110083645,0.00007800481,0.000032978358,0.0004036381,0.00049279124,0.9618036,0.03561393],"study_design_scores_gemma":[0.000037113896,0.00002383297,0.0019720048,0.0006019273,0.00002785299,0.00017256037,0.00020304095,0.00021945336,0.00039096916,0.00089320214,0.99542433,0.00003378856],"about_ca_topic_score_codex":0.006324911,"about_ca_topic_score_gemma":0.013703101,"teacher_disagreement_score":0.989026,"about_ca_system_score_codex":0.00270742,"about_ca_system_score_gemma":0.015301317,"threshold_uncertainty_score":0.8613424},"labels":[],"label_agreement":null},{"id":"W4294585872","doi":"10.7287/peerj.13843v0.1/reviews/3","title":"Peer Review #3 of \"Non-synonymous to synonymous substitutions suggest that orthologs tend to keep their functions, while paralogs are a source of functional novelty (v0.1)\"","year":2022,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Novelty; Evolutionary biology; Biology; Genetics; Psychology; Social psychology","score_opus":0.05879793768669135,"score_gpt":0.30656117654761206,"score_spread":0.24776323886092072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294585872","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007353167,0.013476909,0.016644541,0.12394713,0.46088755,0.0061895344,0.0319426,0.0076186,0.33194],"genre_scores_gemma":[0.035850137,0.015615066,0.01260507,0.013790994,0.07434891,0.0022018577,0.03393862,0.0069678226,0.8046815],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9874374,0.0017983288,0.0013172713,0.0011581191,0.0070325355,0.001256267],"domain_scores_gemma":[0.89948326,0.005265287,0.002511418,0.0054215803,0.07781563,0.009502701],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008619486,0.0013149494,0.002156448,0.0050106407,0.0044395593,0.010573476,0.0035481078,0.0031835337,0.43846136],"category_scores_gemma":[0.07106218,0.0010323796,0.0021024654,0.0035026232,0.0015091052,0.005434531,0.0059030135,0.0023723228,0.33954844],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003599562,0.000012807149,0.0004375876,0.0007120704,0.000021992122,0.0001236205,0.00007550351,0.000033088792,0.00048270402,0.0005861913,0.9605327,0.036945697],"study_design_scores_gemma":[0.000023815765,0.000016197568,0.0014900106,0.00039452055,0.000019005627,0.00016026695,0.00017242772,0.00018322901,0.00035495558,0.0007662808,0.99639577,0.00002346819],"about_ca_topic_score_codex":0.0074446443,"about_ca_topic_score_gemma":0.016510228,"teacher_disagreement_score":0.9913805,"about_ca_system_score_codex":0.0026075616,"about_ca_system_score_gemma":0.014791931,"threshold_uncertainty_score":0.80096674},"labels":[],"label_agreement":null},{"id":"W4294586986","doi":"10.7287/peerj.13843v0.1/reviews/1","title":"Peer Review #1 of \"Non-synonymous to synonymous substitutions suggest that orthologs tend to keep their functions, while paralogs are a source of functional novelty (v0.1)\"","year":2022,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Novelty; Evolutionary biology; Biology; Synonymous substitution; Computational biology; Genetics; Psychology; Gene; Social psychology; Codon usage bias; Genome","score_opus":0.05879793768669135,"score_gpt":0.30656117654761206,"score_spread":0.24776323886092072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294586986","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062915054,0.01561894,0.0141672185,0.13503371,0.5831778,0.0065587633,0.03464099,0.0060904943,0.19842067],"genre_scores_gemma":[0.040360216,0.02371067,0.015507371,0.018774481,0.122703984,0.0033054021,0.053056743,0.0069789495,0.71560234],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9842902,0.0026346582,0.0018242722,0.0015092805,0.008295346,0.0014462863],"domain_scores_gemma":[0.84005183,0.009321956,0.004289018,0.007991663,0.12592393,0.012421679],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.010988897,0.0014574662,0.0027924653,0.0057763536,0.0044805654,0.010775745,0.003976028,0.003732277,0.3762665],"category_scores_gemma":[0.09817415,0.0012198404,0.002254687,0.004105735,0.0016180583,0.0064034965,0.0062848763,0.0027973547,0.29300916],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000045638844,0.000014709396,0.0004787542,0.0010423448,0.000027701084,0.00012237793,0.000082786486,0.000034380588,0.00043400566,0.00051769824,0.9591184,0.0380813],"study_design_scores_gemma":[0.00003495498,0.000023812325,0.002077592,0.00062564656,0.000029093395,0.00018353219,0.00021024256,0.00022092524,0.0003915788,0.0008856853,0.9952832,0.00003373335],"about_ca_topic_score_codex":0.0063404622,"about_ca_topic_score_gemma":0.013974562,"teacher_disagreement_score":0.9890111,"about_ca_system_score_codex":0.0027226703,"about_ca_system_score_gemma":0.0154848965,"threshold_uncertainty_score":0.88968015},"labels":[],"label_agreement":null},{"id":"W4294661678","doi":"","title":"Ontology design for pneumonia diagnostic","year":2018,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"","keywords":"Computer science; Pneumonia; Ontology; Data science; Medicine; Philosophy; Epistemology","score_opus":0.024287571293629604,"score_gpt":0.26371921873729304,"score_spread":0.23943164744366344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294661678","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012257457,0.0019763873,0.96705025,0.0027321498,0.00024054534,0.00051256834,0.0055474658,0.0041030347,0.0055800974],"genre_scores_gemma":[0.18992369,0.002343935,0.7866587,0.00086902565,0.00014738626,0.00047532562,0.014230914,0.00050569547,0.0048452825],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968934,0.0006686754,0.0006134707,0.0006436601,0.00095591927,0.00022486212],"domain_scores_gemma":[0.9962853,0.0013935431,0.00028417297,0.00057821133,0.0012064016,0.00025226854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00343037,0.00070881104,0.0008522203,0.0039376603,0.0011642888,0.003442024,0.0014000715,0.0011339054,0.004028244],"category_scores_gemma":[0.009465899,0.00048491053,0.0024474417,0.002711503,0.00084878854,0.0042513763,0.0026317954,0.0013649688,0.001556301],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004454292,0.00034815277,0.012442417,0.0028804967,0.00048744513,0.001803708,0.0016594193,0.023548221,0.023013135,0.20238072,0.039077286,0.6919136],"study_design_scores_gemma":[0.00009046347,0.00016879605,0.0056349034,0.0011874614,0.00087534456,0.0021169311,0.0012240543,0.2692888,0.035397645,0.4437696,0.24011236,0.0001335871],"about_ca_topic_score_codex":0.009118538,"about_ca_topic_score_gemma":0.01129755,"teacher_disagreement_score":0.009118538,"about_ca_system_score_codex":0.0021521645,"about_ca_system_score_gemma":0.0040215747,"threshold_uncertainty_score":0.018141747},"labels":[],"label_agreement":null},{"id":"W4294916434","doi":"10.2196/41136","title":"Relation Extraction in Biomedical Texts Based on Multi-Head Attention Model With Syntactic Dependency Feature: Modeling Study","year":2022,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Relationship extraction; Computer science; Artificial intelligence; Natural language processing; Sentence; Relation (database); Feature (linguistics); Embedding; Feature extraction; Biomedical text mining; Task (project management); Information extraction; Data mining; Text mining","score_opus":0.02835044199180022,"score_gpt":0.3284713126778678,"score_spread":0.3001208706860676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4294916434","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30812073,0.002418226,0.6800504,0.0016737038,0.0001237371,0.00018250864,0.00086007855,0.0014279324,0.005142758],"genre_scores_gemma":[0.93892264,0.00071917457,0.05230336,0.00025238807,0.00012882058,0.00018142027,0.0009426737,0.00006181044,0.0064877374],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99962723,0.000078600024,0.000024138613,0.00015361235,0.000060090268,0.000056343302],"domain_scores_gemma":[0.99908257,0.0005958171,0.00009621074,0.000043367687,0.00014714144,0.000034946792],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00068993395,0.00075107976,0.0006428402,0.0013585886,0.0004669217,0.0006791451,0.001102786,0.0009179604,0.0021057995],"category_scores_gemma":[0.0018479343,0.0003333922,0.0013853373,0.0012232799,0.0004624982,0.0019266631,0.0007090992,0.0010797431,0.0005003359],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006594024,0.0007263617,0.018381955,0.0004081697,0.0003711775,0.0013297208,0.00087959867,0.58948225,0.018841125,0.024441933,0.009122341,0.335356],"study_design_scores_gemma":[0.0000039442025,0.000017819153,0.0007673979,0.0000028076445,0.000020892548,0.000027815015,0.000011046179,0.996269,0.0005687064,0.0020853295,0.00022030038,0.0000049417254],"about_ca_topic_score_codex":0.019151928,"about_ca_topic_score_gemma":0.02083818,"teacher_disagreement_score":0.019151928,"about_ca_system_score_codex":0.0011626136,"about_ca_system_score_gemma":0.0011072255,"threshold_uncertainty_score":0.03808087},"labels":[],"label_agreement":null},{"id":"W4295894894","doi":"10.1016/s2215-0366(22)00309-1","title":"Applied ontology for phenomenological psychopathology? A cautionary tale – Authors' reply","year":2022,"lang":"en","type":"letter","venue":"The Lancet Psychiatry","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Psychopathology; Ontology; Psychology; Computer science; Data science; Information retrieval; Psychiatry; Epistemology; Philosophy","score_opus":0.032774888075954314,"score_gpt":0.2936450464642955,"score_spread":0.2608701583883412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295894894","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00005293535,0.00037510198,0.00006421851,0.9916563,0.0076692705,0.0000028382349,0.000017472814,0.0000069439648,0.00015488362],"genre_scores_gemma":[0.0010489906,0.0003040725,0.00022007864,0.9803725,0.017295286,0.00002859889,0.000010602844,0.000014913499,0.0007048462],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9859491,0.0046349326,0.0023549919,0.0023635842,0.0034608168,0.0012365691],"domain_scores_gemma":[0.8642844,0.10508931,0.003500126,0.0037027267,0.017404892,0.006018545],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021073107,0.0009614757,0.0028478994,0.0024372612,0.007924255,0.008656286,0.0047330456,0.08187707,0.0064586652],"category_scores_gemma":[0.1357008,0.0013940319,0.0023811653,0.0020543446,0.01583121,0.012598288,0.006910897,0.10122893,0.005073267],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027929415,0.00001454128,0.00029608572,0.000067629204,0.000039774754,0.0004329468,0.0005400758,0.00003182681,0.000040211606,0.0065148403,0.9876583,0.004335765],"study_design_scores_gemma":[0.00026886238,0.00004768168,0.0014902586,0.0015248624,0.00014982224,0.0017008077,0.00351872,0.0009035041,0.00024939445,0.09928828,0.89054877,0.00030893178],"about_ca_topic_score_codex":0.015098502,"about_ca_topic_score_gemma":0.019767947,"teacher_disagreement_score":0.08187707,"about_ca_system_score_codex":0.0071193953,"about_ca_system_score_gemma":0.011731524,"threshold_uncertainty_score":0.11144656},"labels":[],"label_agreement":null},{"id":"W4295916119","doi":"10.25071/1708-6701.40425","title":"CAML members reflect (3)","year":2021,"lang":"fr","type":"article","venue":"CAML Review / Revue de l ACBM","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Computational biology; Biology","score_opus":0.02542008628059269,"score_gpt":0.3199840549321165,"score_spread":0.29456396865152384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295916119","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00054939673,0.002579461,0.0025913564,0.12933874,0.04480939,0.0004237739,0.013826877,0.004610161,0.8012708],"genre_scores_gemma":[0.0019537073,0.0012184734,0.0025586488,0.027464395,0.0054747523,0.00056876533,0.008272694,0.0014695739,0.9510189],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99213964,0.0013750854,0.00036765498,0.00081867125,0.0041952594,0.0011036779],"domain_scores_gemma":[0.9698963,0.0036244462,0.0014946401,0.0033030729,0.012698397,0.0089832],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.010913522,0.0011328722,0.00081469305,0.0039814133,0.003088714,0.011470656,0.0028578287,0.007217277,0.58599395],"category_scores_gemma":[0.040519387,0.0006572019,0.0010793945,0.003959281,0.0014010766,0.0051482227,0.0072146496,0.0038209045,0.59540886],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000008441881,0.0000042488336,0.00006590818,0.000030861982,7.271621e-7,0.0000061312694,0.000024328907,0.0000025447928,0.00004275598,0.0011932809,0.9872163,0.011404411],"study_design_scores_gemma":[0.0000035162197,0.0000021186283,0.00010923502,0.00003734829,0.0000011580398,0.00000744043,0.000029477891,0.000008042699,0.000025361294,0.00029323075,0.9994809,0.0000022414702],"about_ca_topic_score_codex":0.01328916,"about_ca_topic_score_gemma":0.0363327,"teacher_disagreement_score":0.58599395,"about_ca_system_score_codex":0.0040770825,"about_ca_system_score_gemma":0.010954084,"threshold_uncertainty_score":0.5905294},"labels":[],"label_agreement":null},{"id":"W4295942242","doi":"10.25071/1708-6701.40424","title":"CAML members reflect (2)","year":2021,"lang":"fr","type":"article","venue":"CAML Review / Revue de l ACBM","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.02537138990852821,"score_gpt":0.319708148998933,"score_spread":0.2943367590904048,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4295942242","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00044228352,0.0029215107,0.0021840432,0.37723035,0.08835344,0.0003372088,0.006419783,0.0030584342,0.519053],"genre_scores_gemma":[0.0019813902,0.00137915,0.002515032,0.06689071,0.010482838,0.00042118455,0.0037121198,0.0009928469,0.9116248],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9882097,0.001958659,0.00049732515,0.0012447502,0.0066612605,0.001428276],"domain_scores_gemma":[0.9485229,0.0063882265,0.0024614367,0.0046969745,0.023093285,0.014837116],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.016277485,0.001185102,0.00091631233,0.003568564,0.004063639,0.013845673,0.0028525384,0.010685966,0.48483253],"category_scores_gemma":[0.058278233,0.00072940567,0.0011239129,0.0032473844,0.0022032757,0.006442435,0.0073219393,0.006257374,0.48413193],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000006697987,0.0000040859522,0.000052296844,0.000020860676,7.090597e-7,0.0000059374333,0.000021504544,0.000002337552,0.00003333625,0.0013211092,0.9911431,0.0073880185],"study_design_scores_gemma":[0.000003575873,0.0000023512507,0.00009118707,0.00003276854,0.0000013838975,0.000007645841,0.000034979672,0.000008685366,0.000027592088,0.0003319098,0.9994553,0.0000026573455],"about_ca_topic_score_codex":0.016742827,"about_ca_topic_score_gemma":0.047721013,"teacher_disagreement_score":0.5151675,"about_ca_system_score_codex":0.0058847778,"about_ca_system_score_gemma":0.013736188,"threshold_uncertainty_score":0.7348239},"labels":[],"label_agreement":null},{"id":"W4296091000","doi":"10.25071/1708-6701.40422","title":"CAML members reflect (1)","year":2021,"lang":"en","type":"article","venue":"CAML Review / Revue de l ACBM","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Computational biology; Natural language processing; Biology","score_opus":0.02174173544848011,"score_gpt":0.31503607712411896,"score_spread":0.29329434167563884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296091000","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005077305,0.007094938,0.0033194257,0.5176294,0.10486625,0.0002414138,0.0057512475,0.0027553842,0.3578342],"genre_scores_gemma":[0.0045583,0.005172279,0.0053405534,0.08813649,0.01906537,0.0005206022,0.005745566,0.0017008699,0.8697599],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9889352,0.002192145,0.00040116164,0.0012271345,0.0061444407,0.0010999228],"domain_scores_gemma":[0.94269156,0.0070148073,0.0028560509,0.0038381964,0.027814366,0.015784943],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.013512042,0.0011957905,0.0008928155,0.0031648264,0.0030193313,0.012157462,0.002556203,0.0072260113,0.38748008],"category_scores_gemma":[0.06372205,0.0005682454,0.00085272506,0.0028971238,0.0017697285,0.0067818044,0.007154624,0.006157046,0.36427337],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009763003,0.0000041063045,0.000057517485,0.000027449858,8.97325e-7,0.0000068857303,0.000023197203,0.0000030754447,0.000038368657,0.0015122086,0.9865449,0.011771675],"study_design_scores_gemma":[0.0000033332908,0.0000027524588,0.000078565245,0.00005283849,0.0000016237294,0.000011651728,0.00003906502,0.0000119303495,0.00003394743,0.00043953393,0.99932206,0.0000026583598],"about_ca_topic_score_codex":0.011489331,"about_ca_topic_score_gemma":0.028338479,"teacher_disagreement_score":0.38748008,"about_ca_system_score_codex":0.005620095,"about_ca_system_score_gemma":0.014984299,"threshold_uncertainty_score":0.8736853},"labels":[],"label_agreement":null},{"id":"W4297461828","doi":"10.1109/tvcg.2022.3209444","title":"<i>ChartWalk</i>: Navigating large collections of text notes in electronic health records for clinical chart review","year":2022,"lang":"en","type":"article","venue":"IEEE Transactions on Visualization and Computer Graphics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Addiction and Mental Health; University of Toronto","funders":"","keywords":"Computer science; Chart; Health records; Electronic health record; Information retrieval; Data visualization; Data science; Visualization; World Wide Web; Data mining; Health care","score_opus":0.03657597239781057,"score_gpt":0.3798537146484933,"score_spread":0.34327774225068275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4297461828","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022723345,0.0003307711,0.9196552,0.0026131223,0.0002622298,0.002556719,0.0030760288,0.042489868,0.0062927078],"genre_scores_gemma":[0.058306,0.00023064845,0.93146706,0.00051504176,0.00006935382,0.0019121111,0.0017824644,0.00264367,0.003073546],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99304175,0.004491021,0.00058340863,0.00075048546,0.00088575535,0.00024766766],"domain_scores_gemma":[0.9671332,0.02210282,0.0017293962,0.0038318709,0.003651555,0.001551213],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010087919,0.0018191237,0.0005169462,0.0021230308,0.0010363,0.0033773682,0.002252095,0.0017371583,0.011491703],"category_scores_gemma":[0.032265164,0.0007134726,0.0009351644,0.0017317589,0.00136957,0.004391996,0.0037333588,0.0012989956,0.0031892904],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002084704,0.0006826941,0.009074043,0.007122203,0.00025432027,0.0025337166,0.039283272,0.0065383404,0.101965666,0.021363113,0.16486867,0.64422923],"study_design_scores_gemma":[0.0007926759,0.0031157613,0.018791154,0.0037670524,0.0004075741,0.0042453874,0.011625209,0.07602939,0.107413225,0.035891254,0.7370533,0.0008680762],"about_ca_topic_score_codex":0.0010525725,"about_ca_topic_score_gemma":0.0022108308,"teacher_disagreement_score":0.011491703,"about_ca_system_score_codex":0.00073382875,"about_ca_system_score_gemma":0.001921492,"threshold_uncertainty_score":0.053350687},"labels":[],"label_agreement":null},{"id":"W4298842718","doi":"10.6084/m9.figshare.c.6069838","title":"Integration and publication of heterogeneous text-mined relationships on the Semantic Web","year":2022,"lang":"en","type":"other","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"Social Semantic Web; World Wide Web; Semantic Web; Computer science; Semantic Web Stack; Semantic analytics; Information retrieval","score_opus":0.017860925297712788,"score_gpt":0.22835422513849618,"score_spread":0.21049329984078338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4298842718","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26195055,0.0039251214,0.65140474,0.0043698344,0.000642491,0.0006739711,0.025462676,0.015062582,0.03650806],"genre_scores_gemma":[0.4272431,0.0024141667,0.52884114,0.00058678107,0.00020148125,0.00028038325,0.03541548,0.0013965765,0.0036209258],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99440867,0.0019201563,0.00083853264,0.0007715486,0.0018488369,0.00021220259],"domain_scores_gemma":[0.98212576,0.008972314,0.0021105155,0.0035369943,0.0028819088,0.0003726388],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007145745,0.0005295804,0.0007481381,0.012944211,0.00094379456,0.0040336596,0.001014641,0.0010749293,0.001857593],"category_scores_gemma":[0.02109147,0.00039220497,0.0011140009,0.014570352,0.0009344643,0.007137063,0.002634207,0.0009835777,0.0008663602],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007226295,0.0008253664,0.02984724,0.003297682,0.0006807543,0.0072066057,0.0048160646,0.02852078,0.05888135,0.16691384,0.03975435,0.65853333],"study_design_scores_gemma":[0.00012273807,0.00023606794,0.0256882,0.0015221623,0.00082746614,0.0028499665,0.0031075112,0.23230165,0.13395926,0.1776041,0.42153174,0.0002491817],"about_ca_topic_score_codex":0.0019913607,"about_ca_topic_score_gemma":0.0022163992,"teacher_disagreement_score":0.012944211,"about_ca_system_score_codex":0.0010832467,"about_ca_system_score_gemma":0.0020600753,"threshold_uncertainty_score":0.037790775},"labels":[],"label_agreement":null},{"id":"W4299956229","doi":"","title":"Ismb 2002. Proceedings of the 10th International Conference on Intelligent Systems for Molecular Biology. Edmonton, Canada, August 3-7, 2002.","year":2002,"lang":"en","type":"other","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Computer science; Computational biology; Operations research; Biology; Engineering","score_opus":0.026302471726059245,"score_gpt":0.24607150359503754,"score_spread":0.2197690318689783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4299956229","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071411463,0.15731126,0.20085967,0.03063225,0.03245869,0.0013287052,0.12817183,0.06475411,0.3773423],"genre_scores_gemma":[0.0066573294,0.0754907,0.09335446,0.003309477,0.002215474,0.0005462457,0.17249227,0.0071392553,0.6387948],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99915385,0.00018077865,0.00010295104,0.0001461022,0.00035215652,0.00006415688],"domain_scores_gemma":[0.9956943,0.0010789399,0.00015794717,0.00047353617,0.001742761,0.0008524663],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032276837,0.0025450808,0.0024315605,0.0038570191,0.0011409296,0.0054962137,0.0021758832,0.0016803594,0.12610628],"category_scores_gemma":[0.0043810057,0.0011285131,0.0009683048,0.004198831,0.00070572866,0.0044416017,0.0020212894,0.0024890916,0.15835232],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000110381094,0.00006688444,0.000240633,0.0005556676,0.00003976991,0.00008034111,0.000054012908,0.00029098886,0.0011892017,0.0013670941,0.85293484,0.14307019],"study_design_scores_gemma":[0.000029058012,0.00001724843,0.0009612773,0.00027238592,0.000056034027,0.00017836012,0.00007822914,0.001006954,0.0008660332,0.0020856773,0.99442554,0.000023217497],"about_ca_topic_score_codex":0.02582769,"about_ca_topic_score_gemma":0.076471575,"teacher_disagreement_score":0.9741723,"about_ca_system_score_codex":0.001766754,"about_ca_system_score_gemma":0.0044465363,"threshold_uncertainty_score":0.42186755},"labels":[],"label_agreement":null},{"id":"W4302275377","doi":"10.1093/database/baac084","title":"Overview of the COVID-19 text mining tool interactive demonstration track in BioCreative VII","year":2022,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal; Institute for Research in Immunology and Cancer","funders":"U.S. National Library of Medicine; National Institute of General Medical Sciences; National Human Genome Research Institute; Canadian Institutes of Health Research; National Institutes of Health","keywords":"Coronavirus disease 2019 (COVID-19); Computer science; Track (disk drive); Severe acute respiratory syndrome coronavirus 2 (SARS-CoV-2); 2019-20 coronavirus outbreak; Information retrieval; Natural language processing; Artificial intelligence; World Wide Web; Virology; Medicine","score_opus":0.05752538837657176,"score_gpt":0.34534285423608096,"score_spread":0.2878174658595092,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4302275377","genre_codex":"software","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022978513,0.0014360823,0.35889694,0.0014966483,0.0007053954,0.005642165,0.09800353,0.46152675,0.04931394],"genre_scores_gemma":[0.053225115,0.00105182,0.6056722,0.0015344328,0.00028163023,0.008377775,0.21160801,0.045988392,0.07226072],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99741346,0.00046763447,0.0003375694,0.00058853644,0.0010091335,0.00018369887],"domain_scores_gemma":[0.9917549,0.005209488,0.00028000027,0.00079377944,0.0012777062,0.00068422913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048654475,0.0018562152,0.0010910326,0.0042309244,0.0010782768,0.0041075004,0.004466458,0.002000463,0.08958161],"category_scores_gemma":[0.010109628,0.0014260419,0.001485525,0.0022789466,0.00050388306,0.004605745,0.0027993466,0.0019138933,0.04419575],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020786885,0.0015404074,0.006586143,0.003352842,0.0002397862,0.002019366,0.0027269567,0.00338381,0.032061554,0.004220346,0.5974427,0.3443474],"study_design_scores_gemma":[0.0008245967,0.0010771417,0.009266864,0.0007348295,0.00010968974,0.0016073092,0.0005503328,0.04023557,0.024903271,0.0033575161,0.9170151,0.00031779334],"about_ca_topic_score_codex":0.005144983,"about_ca_topic_score_gemma":0.0063827806,"teacher_disagreement_score":0.08958161,"about_ca_system_score_codex":0.0008259601,"about_ca_system_score_gemma":0.0017038633,"threshold_uncertainty_score":0.29968035},"labels":[],"label_agreement":null},{"id":"W4302810357","doi":"10.1145/3487553.3524875","title":"The International Workshop on Semantics-enabled Biomedical Literature Analytics (SeBiLAn)","year":2022,"lang":"en","type":"article","venue":"Companion Proceedings of the Web Conference 2022","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Citation; Commonwealth; Analytics; Library science; Computer science; World Wide Web; Data science; History; Archaeology","score_opus":0.019400183264635856,"score_gpt":0.25924942925375044,"score_spread":0.23984924598911458,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4302810357","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008938194,0.07641731,0.61776567,0.071443595,0.029206125,0.0016657272,0.035852548,0.027272105,0.13143867],"genre_scores_gemma":[0.042351827,0.049564928,0.5891148,0.01679012,0.0059774998,0.0017438976,0.12858307,0.007062052,0.15881185],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9895336,0.004219769,0.0010847406,0.0017289596,0.0028774498,0.000555469],"domain_scores_gemma":[0.9814858,0.008175693,0.00043363322,0.0040224777,0.003820167,0.0020623168],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.016953422,0.0016253485,0.0020289952,0.009006816,0.0018549672,0.010697148,0.0029179703,0.002542019,0.059832934],"category_scores_gemma":[0.028456999,0.0009959157,0.0029622049,0.0077537578,0.0019617288,0.015734946,0.012622922,0.0040187263,0.04117181],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002498324,0.00027601497,0.0007016459,0.001547694,0.0001729288,0.0002972781,0.0008747169,0.0015083738,0.0034258475,0.050786704,0.48462635,0.4555326],"study_design_scores_gemma":[0.00002915295,0.000042117394,0.00066808437,0.0008164168,0.000041331572,0.00025028863,0.00036902935,0.0058786715,0.001357158,0.06381131,0.9267017,0.00003466485],"about_ca_topic_score_codex":0.0036525438,"about_ca_topic_score_gemma":0.0054793605,"teacher_disagreement_score":0.9909932,"about_ca_system_score_codex":0.002330701,"about_ca_system_score_gemma":0.007988713,"threshold_uncertainty_score":0.2001611},"labels":[],"label_agreement":null},{"id":"W4303453972","doi":"10.21203/rs.3.rs-1813123/v1","title":"An Ontology-based approach for Modelling and Querying Alzheimer’s Disease Data","year":2022,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute on Aging; National Institute of Biomedical Imaging and Bioengineering; Canadian Institutes of Health Research; National Institutes of Health; Genentech; IXICO; H. Lundbeck A/S; Servier; Novartis Pharmaceuticals Corporation; Regione Lazio; Eisai; Northern California Institute for Research and Education; Pfizer; Biogen; BioClinica; F. Hoffmann-La Roche; University of Southern California; European Regional Development Fund; Eli Lilly and Company; U.S. Department of Defense; Meso Scale Diagnostics; Alzheimer's Disease Neuroimaging Initiative; Bristol-Myers Squibb; Alzheimer's Association; Foundation for the National Institutes of Health","keywords":"Ontology; Computer science; Disease; Information retrieval; Data science; Medicine; Epistemology; Philosophy; Internal medicine","score_opus":0.3089350076466285,"score_gpt":0.4669794441964446,"score_spread":0.1580444365498161,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4303453972","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006884447,0.00036412178,0.97696406,0.0010926774,0.000076719785,0.00031300166,0.0058070053,0.0063993763,0.0020985669],"genre_scores_gemma":[0.066725746,0.0006373138,0.9158628,0.00043421466,0.000048144513,0.0003313533,0.013165149,0.0005855168,0.0022097281],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9970498,0.00049397396,0.00068074855,0.00046126414,0.0011723957,0.00014182019],"domain_scores_gemma":[0.99698997,0.0013047544,0.00021103898,0.00075774576,0.00056468556,0.00017180337],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004237243,0.0008955235,0.001106048,0.004672408,0.0013478397,0.00580062,0.0025555682,0.0017614496,0.0019784707],"category_scores_gemma":[0.008270092,0.00081193977,0.0033312473,0.005718609,0.0009045678,0.0066362014,0.003536592,0.002473856,0.0009400267],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066686457,0.0009020457,0.008242174,0.0020416467,0.0011956241,0.0021963026,0.0037925928,0.08401729,0.026896797,0.33630276,0.049093507,0.4846524],"study_design_scores_gemma":[0.00013219858,0.000091752394,0.0025433626,0.0004382228,0.00060678413,0.00087131024,0.0011376047,0.50465393,0.014949349,0.31698018,0.15745158,0.00014368244],"about_ca_topic_score_codex":0.022110362,"about_ca_topic_score_gemma":0.032491107,"teacher_disagreement_score":0.022110362,"about_ca_system_score_codex":0.0019158188,"about_ca_system_score_gemma":0.0031583868,"threshold_uncertainty_score":0.043963373},"labels":[],"label_agreement":null},{"id":"W4308835006","doi":"10.1002/pds.5555","title":"More extreme duplication in FDA Adverse Event Reporting System detected by literature reference normalization and fuzzy string matching","year":2022,"lang":"en","type":"article","venue":"Pharmacoepidemiology and Drug Safety","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Adverse Event Reporting System; Computer science; Data mining; Information retrieval; Levenshtein distance; Normalization (sociology); Medicine; Matching (statistics); String metric; String searching algorithm; Adverse effect; Artificial intelligence; Pattern matching; Internal medicine","score_opus":0.025662152849074908,"score_gpt":0.31312907562108255,"score_spread":0.28746692277200764,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4308835006","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7559217,0.00554836,0.20547622,0.002503789,0.00070456986,0.0009666324,0.010593215,0.0054664207,0.012819013],"genre_scores_gemma":[0.7243852,0.0012085814,0.25283566,0.0007239931,0.00019941489,0.0005044824,0.01675311,0.00037430678,0.0030153466],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9753092,0.0042955796,0.005905763,0.0043306244,0.009508662,0.00065018894],"domain_scores_gemma":[0.9050512,0.041329347,0.01922499,0.012390086,0.021202547,0.0008018304],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019545957,0.00069128384,0.0010139351,0.019516272,0.0015705305,0.00329939,0.0020266108,0.0012724967,0.0021652768],"category_scores_gemma":[0.092702374,0.0004388493,0.0014601703,0.018142112,0.00090104743,0.0027364346,0.002362862,0.0009137431,0.0011315638],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011293376,0.0004264558,0.3456759,0.0027941621,0.0008773555,0.005701009,0.004801033,0.01757329,0.042846777,0.011384601,0.015662828,0.55112725],"study_design_scores_gemma":[0.00019046615,0.0011184713,0.468954,0.0017805145,0.0020644385,0.01070465,0.0061243903,0.18010384,0.18206535,0.026142536,0.120181866,0.0005694436],"about_ca_topic_score_codex":0.007909824,"about_ca_topic_score_gemma":0.006116641,"teacher_disagreement_score":0.980454,"about_ca_system_score_codex":0.002077564,"about_ca_system_score_gemma":0.0038267316,"threshold_uncertainty_score":0.10337013},"labels":[],"label_agreement":null},{"id":"W4309688679","doi":"10.3390/app122211839","title":"Development of an Ontology-Based Solution to Reduce the Spread of Viruses","year":2022,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"New Brunswick Innovation Foundation","keywords":"Ontology; Computer science; Population; Event (particle physics); Social distance; Context (archaeology); Computer security; Coronavirus disease 2019 (COVID-19); Knowledge management; World Wide Web; Geography; Medicine; Environmental health","score_opus":0.051168020991230315,"score_gpt":0.31921406442727035,"score_spread":0.26804604343604005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4309688679","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017603204,0.0003377815,0.96622133,0.001784264,0.00018534866,0.0006495414,0.0008659742,0.004177385,0.00817523],"genre_scores_gemma":[0.07821593,0.0005893334,0.91211236,0.00031266085,0.00003346857,0.00029891325,0.0030839103,0.00022594292,0.0051275324],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984884,0.0002542042,0.00021164824,0.00026718815,0.0006184375,0.00016000275],"domain_scores_gemma":[0.99860555,0.00028892432,0.00016319276,0.00029710925,0.00054252474,0.0001027265],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018371737,0.00060902565,0.0005282339,0.0021390233,0.0014191511,0.002384305,0.0016918886,0.0012435989,0.002116094],"category_scores_gemma":[0.0038815236,0.0003757843,0.0019423484,0.0014511739,0.00045854386,0.0037738301,0.003143142,0.0016140414,0.0008950974],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024834074,0.0013207524,0.0076541626,0.0015355612,0.0004942408,0.0016143583,0.0030333719,0.033887085,0.052686922,0.09418108,0.030284699,0.77305937],"study_design_scores_gemma":[0.000120617806,0.00021517035,0.0063277893,0.000535901,0.0007134129,0.0016880227,0.0051444047,0.48911968,0.049298655,0.09380465,0.35286522,0.00016650974],"about_ca_topic_score_codex":0.0105619645,"about_ca_topic_score_gemma":0.010710094,"teacher_disagreement_score":0.0105619645,"about_ca_system_score_codex":0.00081705954,"about_ca_system_score_gemma":0.0039021852,"threshold_uncertainty_score":0.021000981},"labels":[],"label_agreement":null},{"id":"W4310065946","doi":"10.1016/j.ijmedinf.2022.104928","title":"Towards semantic-driven boolean query formalization for biomedical systematic literature reviews","year":2022,"lang":"en","type":"review","venue":"International Journal of Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Computer science; Query expansion; Unified Medical Language System; Information retrieval; Exploit; Identification (biology); Cluster analysis; Query language; Benchmark (surveying); Query optimization; Web search query; Language model; Data mining; Natural language processing; Artificial intelligence; Search engine","score_opus":0.05587407294549552,"score_gpt":0.39093805143315874,"score_spread":0.3350639784876632,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4310065946","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004274906,0.08162251,0.8859734,0.006874705,0.0005240002,0.0022923716,0.0099386005,0.0042386716,0.004260843],"genre_scores_gemma":[0.038890067,0.0407241,0.899433,0.0019677004,0.00028122574,0.00200504,0.014803666,0.00030790785,0.0015872988],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9728133,0.010121167,0.008348375,0.0019578869,0.006250673,0.0005085717],"domain_scores_gemma":[0.9563665,0.030466232,0.0036187049,0.0033171398,0.0057291077,0.00050230365],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.027759908,0.0019592748,0.0028587666,0.016227458,0.0010008505,0.007600789,0.0027541942,0.0016675792,0.0036405001],"category_scores_gemma":[0.048925743,0.001137469,0.005975629,0.011972988,0.0018143308,0.008116643,0.0057981904,0.0028557302,0.0014121172],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026760693,0.00013547,0.0019093152,0.044037793,0.0014310183,0.00038207739,0.0014815843,0.013582913,0.006115935,0.2398393,0.020531878,0.67028517],"study_design_scores_gemma":[0.0003019849,0.00020078123,0.002547003,0.026959077,0.0032228145,0.0015337712,0.0014519169,0.08808896,0.011393857,0.44254944,0.42143482,0.0003155462],"about_ca_topic_score_codex":0.0054467944,"about_ca_topic_score_gemma":0.009046758,"teacher_disagreement_score":0.9722401,"about_ca_system_score_codex":0.0037061577,"about_ca_system_score_gemma":0.012660711,"threshold_uncertainty_score":0.14681017},"labels":[],"label_agreement":null},{"id":"W4310807794","doi":"10.21742/ijhit.2653-309x.2022.2.1.01","title":"Extending the Power of Problem Oriented Medical Record with Disease Association Discovery: The Case Study of Empowering QL4POMR with OpenTargets","year":2022,"lang":"en","type":"article","venue":"International Journal of Hybrid Innovation Technologies","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"Natural Sciences and Engineering Research Council of Canada; Mitacs","keywords":"Health informatics; Association (psychology); Informatics; Data science; Evidence-based medicine; Computer science; Medical record; Medicine; Psychology; Alternative medicine; Public health; Engineering; Pathology","score_opus":0.009372585970853956,"score_gpt":0.2826924755807543,"score_spread":0.27331988960990033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4310807794","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09159648,0.0020606434,0.8257367,0.041424554,0.0004102277,0.0013568928,0.005465625,0.013791767,0.018157098],"genre_scores_gemma":[0.28405586,0.0013013686,0.7029108,0.0035274427,0.00018584017,0.00047237138,0.0034840251,0.00097503053,0.0030872568],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9721921,0.016722806,0.0034028415,0.0023762211,0.004723017,0.00058299623],"domain_scores_gemma":[0.85608107,0.110279545,0.0045672324,0.021405952,0.005985203,0.0016810324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.035978302,0.000623396,0.00086517644,0.0030745454,0.0013471639,0.006254712,0.0027920203,0.0024882806,0.003958888],"category_scores_gemma":[0.09292169,0.0006118594,0.0016675466,0.004701515,0.0029954966,0.011298296,0.012525178,0.0035313915,0.0016210582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015604646,0.0009054051,0.038322855,0.004158261,0.00033670128,0.0073818793,0.026890572,0.014690546,0.008902701,0.11694045,0.028757522,0.75115263],"study_design_scores_gemma":[0.0005616863,0.0010651719,0.010956375,0.0022854707,0.0005081532,0.011114709,0.008453259,0.12219218,0.028849745,0.20000114,0.6136251,0.00038702303],"about_ca_topic_score_codex":0.003370959,"about_ca_topic_score_gemma":0.0038934834,"teacher_disagreement_score":0.035978302,"about_ca_system_score_codex":0.0020951908,"about_ca_system_score_gemma":0.004214365,"threshold_uncertainty_score":0.1902737},"labels":[],"label_agreement":null},{"id":"W4312064948","doi":"10.1111/hir.12471","title":"Application of text mining to the development and validation of a geographic search filter to facilitate evidence retrieval in Ovid <scp>MEDLINE</scp>: An example from the United States","year":2022,"lang":"en","type":"article","venue":"Health Information & Libraries Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Broadcom (Canada); Vancouver Coastal Health","funders":"","keywords":"MEDLINE; Filter (signal processing); Information retrieval; Computer science; Set (abstract data type); Vocabulary; Identification (biology); Controlled vocabulary; Data science; Data mining; Political science","score_opus":0.10598035444713076,"score_gpt":0.3062490695496477,"score_spread":0.20026871510251693,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312064948","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.323604,0.014765722,0.58009297,0.0155156795,0.00083840196,0.023912407,0.017888024,0.006334538,0.017048292],"genre_scores_gemma":[0.1971762,0.0022762315,0.78672904,0.0010589063,0.00013564447,0.006079741,0.005170693,0.00018900851,0.0011844466],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9596475,0.022496449,0.009857692,0.0021412421,0.005420389,0.00043678025],"domain_scores_gemma":[0.70203656,0.24428995,0.013536776,0.005613324,0.033731237,0.000792265],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.066870056,0.0008523177,0.0019655807,0.015601551,0.0016433758,0.004726579,0.0016244233,0.0015510819,0.0021542872],"category_scores_gemma":[0.23711862,0.00048713267,0.0023073896,0.011940417,0.0008217476,0.0028986407,0.0019248668,0.0008916253,0.00085581705],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018927294,0.0005799874,0.05655416,0.014937515,0.0014028435,0.0013326835,0.005598986,0.010542813,0.013538066,0.007886568,0.020992935,0.86474067],"study_design_scores_gemma":[0.0040269904,0.004599779,0.17215757,0.027409537,0.0073371064,0.0045413235,0.011180733,0.37815255,0.08500563,0.058518123,0.24609514,0.00097559363],"about_ca_topic_score_codex":0.008990058,"about_ca_topic_score_gemma":0.011648931,"teacher_disagreement_score":0.93312997,"about_ca_system_score_codex":0.0031650325,"about_ca_system_score_gemma":0.0086595435,"threshold_uncertainty_score":0.35364687},"labels":[],"label_agreement":null},{"id":"W4312679697","doi":"10.2196/43750","title":"Systematized Nomenclature of Medicine–Clinical Terminology (SNOMED CT) Clinical Use Cases in the Context of Electronic Health Record Systems: Systematic Literature Review","year":2022,"lang":"en","type":"review","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":43,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"SNOMED CT; Terminology; Systematized Nomenclature of Medicine; Context (archaeology); Documentation; Medicine; Protocol (science); Health informatics; Clinical decision support system; Medical physics; Computer science; Artificial intelligence; Pathology; Decision support system; Alternative medicine; Public health","score_opus":0.10878285811022585,"score_gpt":0.44393504622138186,"score_spread":0.335152188111156,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312679697","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00135912,0.9956043,0.00048376116,0.0004951284,0.00019614135,0.00080609106,0.00044640328,0.000013289224,0.0005957811],"genre_scores_gemma":[0.016602458,0.97702163,0.0029685847,0.00078038406,0.00015347195,0.001884044,0.00041908168,0.000011741292,0.00015859213],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9655928,0.013613396,0.01243995,0.001722084,0.006123578,0.00050818],"domain_scores_gemma":[0.8796892,0.08897421,0.020078804,0.0018037631,0.008810765,0.0006433245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02008344,0.0016925354,0.0052500013,0.03185884,0.0012944408,0.003932709,0.0027380276,0.0027129296,0.0044647288],"category_scores_gemma":[0.09482623,0.0011062747,0.0053010224,0.02706467,0.0021155856,0.0061506475,0.0032060319,0.0017266629,0.0005029452],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006696295,0.000018776578,0.0008317167,0.94621325,0.0020351852,0.00020156964,0.0008543001,0.00009876965,0.00014882343,0.0008771361,0.002159804,0.04649357],"study_design_scores_gemma":[0.00004159027,0.0000666764,0.0017736063,0.9654614,0.009449744,0.0005434162,0.00081885094,0.00007690708,0.00012576878,0.0004796779,0.021128159,0.000034186043],"about_ca_topic_score_codex":0.0074060103,"about_ca_topic_score_gemma":0.021027248,"teacher_disagreement_score":0.03185884,"about_ca_system_score_codex":0.0065883547,"about_ca_system_score_gemma":0.026763955,"threshold_uncertainty_score":0.106212676},"labels":[],"label_agreement":null},{"id":"W4313575350","doi":"10.2196/44547","title":"An Ontology-Based Approach for Consolidating Patient Data Standardized With European Norm/International Organization for Standardization 13606 (EN/ISO 13606) Into Joint Observational Medical Outcomes Partnership (OMOP) Repositories: Description of a Methodology","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Salud Carlos III; Barcelona Supercomputing Center; European Commission","keywords":"Computer science; Standardization; Ontology; Information retrieval; Data mining; Data science","score_opus":0.13795739351624942,"score_gpt":0.38102018495190687,"score_spread":0.24306279143565745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4313575350","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021399935,0.000084207364,0.9923152,0.00044468656,0.0000458632,0.0009494873,0.0006552162,0.0021468187,0.001218554],"genre_scores_gemma":[0.011623256,0.00011739983,0.98435175,0.00010458979,0.000018689321,0.00062048447,0.0020368095,0.0002334289,0.0008935467],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98338634,0.0042868014,0.003373186,0.0029581538,0.005446685,0.0005488738],"domain_scores_gemma":[0.9868176,0.003729086,0.0012323353,0.0037779745,0.0037677705,0.0006750854],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01968907,0.0012933846,0.0009883427,0.008351801,0.0022362957,0.008437752,0.0036601906,0.0017564612,0.0027816272],"category_scores_gemma":[0.025031814,0.0011897711,0.0038630753,0.0071284287,0.0019649186,0.00871383,0.008299369,0.0029564183,0.0015478374],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029021106,0.0014004163,0.012140706,0.0018462462,0.0007875107,0.0015814332,0.0072807493,0.022169616,0.02487766,0.23558515,0.025345547,0.66669476],"study_design_scores_gemma":[0.0002936132,0.00045209026,0.009836002,0.0015485695,0.0008351185,0.0022755477,0.005791001,0.40480757,0.047023296,0.21252592,0.3140664,0.0005448893],"about_ca_topic_score_codex":0.013660974,"about_ca_topic_score_gemma":0.016046869,"teacher_disagreement_score":0.01968907,"about_ca_system_score_codex":0.0034135396,"about_ca_system_score_gemma":0.013651255,"threshold_uncertainty_score":0.10412699},"labels":[],"label_agreement":null},{"id":"W4317568373","doi":"10.1093/zoolinnean/zlac107","title":"Renaming taxa on ethical grounds threatens nomenclatural stability and scientific communication","year":2023,"lang":"en","type":"article","venue":"Zoological Journal of the Linnean Society","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":57,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"","keywords":"Biology; Taxon; Zoology; Evolutionary biology; Ecology","score_opus":0.04734969107126128,"score_gpt":0.2943280825544965,"score_spread":0.24697839148323522,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4317568373","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.17554753,0.0029780958,0.18995477,0.4207915,0.019059867,0.00057462,0.0004913945,0.0009648398,0.18963756],"genre_scores_gemma":[0.8202537,0.0011641094,0.0937934,0.054887127,0.0034000624,0.00052050274,0.00039217513,0.00082069996,0.024768272],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8401293,0.09162371,0.015126113,0.010929501,0.038076628,0.00411473],"domain_scores_gemma":[0.5461293,0.26777565,0.034065656,0.09047083,0.05026234,0.011296223],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.16113412,0.0005389441,0.001086777,0.0047778995,0.009744143,0.014494797,0.0036501172,0.0075998465,0.0077050515],"category_scores_gemma":[0.34401628,0.0009230905,0.0008805802,0.0035621247,0.022356259,0.014627121,0.012728208,0.015702588,0.0022070727],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022848774,0.000116744864,0.0112922,0.00039391062,0.00015464691,0.000916282,0.05103996,0.0010195713,0.0055530146,0.68017703,0.03524675,0.21386139],"study_design_scores_gemma":[0.00004919884,0.00006153932,0.0052870535,0.0008225273,0.00009484455,0.0006393115,0.014698855,0.0022230116,0.0026061244,0.63428956,0.33909768,0.00013027416],"about_ca_topic_score_codex":0.0036840003,"about_ca_topic_score_gemma":0.0076876767,"teacher_disagreement_score":0.99025583,"about_ca_system_score_codex":0.005805366,"about_ca_system_score_gemma":0.015721586,"threshold_uncertainty_score":0.8521689},"labels":[],"label_agreement":null},{"id":"W4323241317","doi":"10.5220/0011925700003414","title":"Knowledge Graph Based Trustworthy Medical Code Recommendations","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"","keywords":"Trustworthiness; Computer science; Graph; Code (set theory); Computer security; Programming language; Theoretical computer science","score_opus":0.03818684505989317,"score_gpt":0.3446108730047146,"score_spread":0.30642402794482143,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323241317","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18054864,0.0042101587,0.76976055,0.0068536783,0.0008299043,0.0014371208,0.0152307665,0.005793586,0.015335635],"genre_scores_gemma":[0.72129816,0.0009827715,0.26089776,0.00069359277,0.0002635225,0.00026688576,0.010274684,0.00020807832,0.0051144813],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99575657,0.0007386983,0.00037934163,0.0007594027,0.0021486166,0.00021728587],"domain_scores_gemma":[0.9870313,0.007270537,0.0008571709,0.0012892189,0.0032416016,0.00031016057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017646638,0.00070587324,0.00097246474,0.0071089244,0.0010600125,0.0018838,0.001311092,0.0017129472,0.004061337],"category_scores_gemma":[0.020339422,0.00034590767,0.001090498,0.0036586141,0.00047436022,0.0026670357,0.0014546105,0.0011689517,0.0012060996],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001881146,0.0008673409,0.043971773,0.0013795424,0.0010132199,0.0020663203,0.0009542993,0.10342648,0.014133087,0.02807265,0.06233034,0.73990375],"study_design_scores_gemma":[0.00011807166,0.00018454963,0.0067708935,0.00031590567,0.00052303675,0.000680556,0.0004663158,0.92263526,0.0076897773,0.044715162,0.015852297,0.000048259637],"about_ca_topic_score_codex":0.017992208,"about_ca_topic_score_gemma":0.03564696,"teacher_disagreement_score":0.017992208,"about_ca_system_score_codex":0.0012295112,"about_ca_system_score_gemma":0.0027329403,"threshold_uncertainty_score":0.035774946},"labels":[],"label_agreement":null},{"id":"W4323519786","doi":"10.1093/database/baad005","title":"Chemical identification and indexing in full-text articles: an overview of the NLM-Chem track at BioCreative VII","year":2023,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"H2020 Marie Skłodowska-Curie Actions; U.S. National Library of Medicine; Fundação para a Ciência e a Tecnologia; Natural Sciences and Engineering Research Council of Canada; European Commission; University of Oxford; Albaha University; National Institutes of Health; Nvidia","keywords":"Computer science; Search engine indexing; Named-entity recognition; Identification (biology); Information retrieval; Natural language processing; Task (project management); Normalization (sociology); Artificial intelligence","score_opus":0.06945455733307294,"score_gpt":0.3394277437636137,"score_spread":0.26997318643054075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323519786","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039046157,0.042066477,0.2854689,0.01134251,0.009719356,0.009462548,0.26817256,0.27306092,0.0616606],"genre_scores_gemma":[0.025926676,0.008824418,0.341594,0.002278804,0.0019244646,0.003363805,0.5512696,0.017660294,0.04715796],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97117525,0.005458975,0.004570467,0.0053793103,0.011975932,0.0014400429],"domain_scores_gemma":[0.92494494,0.019710723,0.0053402185,0.012188633,0.030755932,0.0070595494],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.041788716,0.0034959638,0.0034113163,0.029708823,0.0037559094,0.013816978,0.0059196856,0.0038300431,0.03323917],"category_scores_gemma":[0.049935557,0.00226855,0.0042811627,0.01986606,0.0011761334,0.015032706,0.008765157,0.0033885916,0.070258155],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00083313725,0.0008148982,0.006158555,0.0052013854,0.00054337754,0.0004501589,0.0008034177,0.0025312952,0.019172555,0.0019748975,0.5349363,0.42658013],"study_design_scores_gemma":[0.00040050177,0.001247764,0.01471431,0.0012761527,0.00026939483,0.0015167107,0.00055565895,0.026843755,0.038201995,0.0036552139,0.910898,0.00042057593],"about_ca_topic_score_codex":0.0060010552,"about_ca_topic_score_gemma":0.009404777,"teacher_disagreement_score":0.9582113,"about_ca_system_score_codex":0.003936198,"about_ca_system_score_gemma":0.008512912,"threshold_uncertainty_score":0.22100246},"labels":[],"label_agreement":null},{"id":"W4323539565","doi":"10.21203/rs.3.rs-2640617/v1","title":"Cerebrovascular disease case identification in inpatient electronic medical record data using natural language processing","year":2023,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Alberta Health Services; University of Calgary","funders":"Canadian Institutes of Health Research","keywords":"Chart; Medical record; Medicine; Electronic medical record; Electronic health record; Coding (social sciences); Artificial intelligence; Predictive value; Clinical decision support system; Identification (biology); Medical diagnosis; Computer science; Natural language processing; Machine learning; Medical emergency; Statistics; Internal medicine; Decision support system; Pathology; Health care","score_opus":0.12247984123055952,"score_gpt":0.4628552906587924,"score_spread":0.34037544942823283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323539565","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.874867,0.0013753976,0.09503268,0.0020046968,0.0000798658,0.0018991603,0.021037987,0.0012162331,0.0024869393],"genre_scores_gemma":[0.8433693,0.00037048978,0.13908197,0.00038794018,0.00010084435,0.0008146117,0.015614168,0.000019710302,0.0002410278],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9960318,0.0018518767,0.0006513538,0.0007647177,0.00057259476,0.00012759244],"domain_scores_gemma":[0.9802325,0.014337807,0.0029592987,0.00060799497,0.0016777188,0.00018465902],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004832333,0.00045563962,0.00044426284,0.005155134,0.0003476993,0.001176185,0.0007564768,0.00042373224,0.00075475295],"category_scores_gemma":[0.019722909,0.00018435973,0.00074068026,0.0025799407,0.00039802125,0.00092628115,0.0007032881,0.00052867085,0.0002788821],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050794857,0.0009142516,0.7105111,0.0016285665,0.0003386766,0.0012943309,0.001341831,0.026791312,0.004752742,0.0012978988,0.008510668,0.2421106],"study_design_scores_gemma":[0.00013940064,0.0004324272,0.36036423,0.00065043656,0.00028330978,0.0011197021,0.0019776202,0.6151593,0.006247575,0.0065110996,0.0069991467,0.00011586337],"about_ca_topic_score_codex":0.012597587,"about_ca_topic_score_gemma":0.016231894,"teacher_disagreement_score":0.012597587,"about_ca_system_score_codex":0.0012893535,"about_ca_system_score_gemma":0.0018927222,"threshold_uncertainty_score":0.025556087},"labels":[],"label_agreement":null},{"id":"W4323844704","doi":"10.1021/acsenergylett.3c00300","title":"Mastering the Art of Scientific Publication – Part 2","year":2023,"lang":"en","type":"article","venue":"ACS Energy Letters","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Nanotechnology; Engineering physics; Engineering ethics; Chemistry; Engineering; Materials science","score_opus":0.019818524384737773,"score_gpt":0.23807836422780068,"score_spread":0.2182598398430629,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4323844704","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00067014626,0.009370023,0.00661427,0.017750988,0.06509764,0.00071365817,0.0017162911,0.0028045082,0.89526254],"genre_scores_gemma":[0.0033595292,0.009995765,0.0032373834,0.004361576,0.01262959,0.00037063865,0.0014349767,0.0013461905,0.96326435],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99093467,0.000766031,0.00060568977,0.00092511013,0.0058442033,0.00092431915],"domain_scores_gemma":[0.9808292,0.0025670484,0.001024698,0.0035448954,0.008107366,0.0039267843],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006250209,0.0016011404,0.0013849753,0.0035521819,0.0028359257,0.017399821,0.0023747706,0.0047081714,0.70388126],"category_scores_gemma":[0.022505667,0.0012193286,0.0019727768,0.0032658656,0.001987604,0.008254698,0.0048126797,0.0051084575,0.6832179],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000030434803,0.00003613378,0.00010598747,0.00037053286,0.000006922484,0.00005513691,0.000049482416,0.000060158476,0.0011785256,0.007023356,0.87842757,0.11265574],"study_design_scores_gemma":[0.0000030889728,0.000012707712,0.0001417028,0.000085013286,0.0000018516214,0.00006891076,0.000018353618,0.000032196127,0.00018793154,0.0009904108,0.9984522,0.000005694867],"about_ca_topic_score_codex":0.0011969934,"about_ca_topic_score_gemma":0.002079759,"teacher_disagreement_score":0.9937498,"about_ca_system_score_codex":0.0035230352,"about_ca_system_score_gemma":0.0068855835,"threshold_uncertainty_score":0.4223774},"labels":[],"label_agreement":null},{"id":"W4328107868","doi":"10.1016/j.isci.2023.106460","title":"Biomedical discovery through the integrative biomedical knowledge hub (iBKH)","year":2023,"lang":"en","type":"article","venue":"iScience","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute; HEC Montréal","funders":"National Center for Complementary and Integrative Health; National Institute on Aging; National Institutes of Health; National Science Foundation","keywords":"Biomedicine; Computer science; Repurposing; Scalability; Data science; Knowledge extraction; Cheminformatics; Drug discovery; Artificial intelligence; Bioinformatics; Engineering","score_opus":0.03057684809467305,"score_gpt":0.3351023296109865,"score_spread":0.3045254815163135,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4328107868","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046384677,0.0031930585,0.89126325,0.0024721506,0.00026250046,0.0006023124,0.01672985,0.022855077,0.016237192],"genre_scores_gemma":[0.29298747,0.002510125,0.6751998,0.0005127946,0.00012429913,0.00027352187,0.022893507,0.0010434361,0.0044550793],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999033,0.00023462958,0.00007791122,0.0002869636,0.00028091637,0.00008666694],"domain_scores_gemma":[0.99696594,0.0012592268,0.00028615035,0.00069975556,0.0005057228,0.0002831602],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002548452,0.000614836,0.00066377164,0.008134648,0.0009186999,0.0023177708,0.0012977121,0.00071797706,0.0051815165],"category_scores_gemma":[0.006184684,0.00047374074,0.0010282495,0.00596182,0.0009006427,0.003994129,0.0043451516,0.0010321003,0.0015215353],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00070141646,0.00034304083,0.01612465,0.0036843906,0.00075085234,0.0013939266,0.0018196647,0.035469476,0.03225683,0.20688537,0.06422204,0.63634837],"study_design_scores_gemma":[0.00020103768,0.0002282905,0.012238164,0.00068918534,0.0008826311,0.0012898408,0.0011256776,0.24026768,0.04647729,0.39182642,0.30457613,0.00019761994],"about_ca_topic_score_codex":0.0038199334,"about_ca_topic_score_gemma":0.0054592956,"teacher_disagreement_score":0.008134648,"about_ca_system_score_codex":0.0009036853,"about_ca_system_score_gemma":0.0028044772,"threshold_uncertainty_score":0.017333925},"labels":[],"label_agreement":null},{"id":"W4361853034","doi":"10.2196/46127","title":"A SNOMED CT Mapping Guideline for the Local Terms Used to Document Clinical Findings and Procedures in Electronic Medical Records in South Korea: Methodological Study","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Ministry of Science and ICT, South Korea; Korea Health Industry Development Institute","keywords":"SNOMED CT; Systematized Nomenclature of Medicine; Context (archaeology); Guideline; Interoperability; Computer science; Information retrieval; Unified Medical Language System; Medicine; Terminology; Medical physics; Data mining; World Wide Web; Pathology; Geography","score_opus":0.09416665320763871,"score_gpt":0.4325761629273854,"score_spread":0.33840950971974665,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4361853034","genre_codex":"protocol","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3123348,0.051611144,0.17952023,0.007346877,0.00080729136,0.41093665,0.01568461,0.00036351974,0.021394867],"genre_scores_gemma":[0.21828385,0.01991794,0.49434456,0.0028467057,0.00010512015,0.25275716,0.009950563,0.00011495102,0.0016791802],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9086933,0.045548465,0.0292799,0.0045754053,0.011000422,0.0009024302],"domain_scores_gemma":[0.88261557,0.051519673,0.016792959,0.006600172,0.041371044,0.0011004936],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.12284311,0.00087145664,0.0015094639,0.021951612,0.002497652,0.0036468757,0.0023216512,0.0015212079,0.0027375668],"category_scores_gemma":[0.16254568,0.0010947958,0.003294335,0.019038169,0.0022202418,0.0055277967,0.005205613,0.0017901355,0.0006205039],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062249077,0.002198853,0.08807905,0.1410804,0.0012457487,0.002238473,0.092244364,0.0021329888,0.004796784,0.016424978,0.024604438,0.62433153],"study_design_scores_gemma":[0.0018407182,0.0035373166,0.14168824,0.25545147,0.013315559,0.0048184087,0.22951335,0.007347042,0.0126062725,0.01690412,0.31220257,0.0007748819],"about_ca_topic_score_codex":0.011800853,"about_ca_topic_score_gemma":0.034303036,"teacher_disagreement_score":0.12284311,"about_ca_system_score_codex":0.009063089,"about_ca_system_score_gemma":0.049832534,"threshold_uncertainty_score":0.6496642},"labels":[],"label_agreement":null},{"id":"W4365515241","doi":"10.1038/s41597-023-02106-1","title":"Open science and data sharing in cognitive neuroscience with MouseBytes and MouseBytes+","year":2023,"lang":"en","type":"article","venue":"Scientific Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Robarts Clinical Trials; Western University","funders":"Canadian Open Neuroscience Platform; Canada First Research Excellence Fund","keywords":"Open science; Cognitive neuroscience; Neuroinformatics; Cognition; Neuroscience; Open data; Cognitive science; Data sharing; Psychology; Data science; Computer science; World Wide Web; Medicine; Physics","score_opus":0.19230217751921452,"score_gpt":0.4030885047671502,"score_spread":0.21078632724793567,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4365515241","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010367807,0.0014277534,0.55126894,0.003147505,0.0008364962,0.0010770535,0.06410025,0.35413587,0.01363824],"genre_scores_gemma":[0.09099859,0.0029393295,0.5203103,0.0031296432,0.00051381084,0.00508241,0.2778685,0.082971394,0.016186118],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99383485,0.0012322621,0.0010549347,0.0017433852,0.0017037404,0.00043082525],"domain_scores_gemma":[0.97744286,0.0045555294,0.001749983,0.011579038,0.0020624362,0.0026101393],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.016910909,0.0022752956,0.0020376977,0.005241521,0.0014463121,0.007446328,0.0068042953,0.0020424188,0.020336624],"category_scores_gemma":[0.027371056,0.0019261378,0.003032707,0.0055854977,0.0025993711,0.012200481,0.017566726,0.0039807195,0.018097814],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0064798063,0.0007346098,0.012650665,0.0046956013,0.0024003158,0.0019857858,0.003151099,0.009516703,0.04551072,0.14311332,0.55788326,0.2118782],"study_design_scores_gemma":[0.0010381939,0.00048055645,0.008110806,0.0008440759,0.000505922,0.0012698898,0.00049898075,0.031739753,0.041092295,0.15951729,0.7543618,0.00054047833],"about_ca_topic_score_codex":0.002424559,"about_ca_topic_score_gemma":0.003019818,"teacher_disagreement_score":0.9931957,"about_ca_system_score_codex":0.001233389,"about_ca_system_score_gemma":0.0047212476,"threshold_uncertainty_score":0.089434505},"labels":[],"label_agreement":null},{"id":"W4366078185","doi":"10.5210/disco.v5i0.2680","title":"MLTrends: Graphing MEDLINE term usage over time","year":2010,"lang":"en","type":"article","venue":"Journal of Biomedical Discovery and Collaboration","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital","funders":"","keywords":"MEDLINE; Computer science; Term (time); Information retrieval; Web of science","score_opus":0.003948985190073336,"score_gpt":0.2569139647097169,"score_spread":0.25296497951964353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366078185","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.11398057,0.00313663,0.17493309,0.001613979,0.00047524142,0.0010538048,0.45418024,0.2348417,0.015784793],"genre_scores_gemma":[0.21958987,0.004546588,0.4513044,0.00037782814,0.00027285118,0.0019664124,0.29770043,0.015280259,0.008961459],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984906,0.0003007641,0.00035998088,0.00032521284,0.00044858447,0.000074886964],"domain_scores_gemma":[0.98864913,0.0072134156,0.0016086113,0.00090765196,0.0013768186,0.00024439118],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0020707292,0.001176302,0.0007545268,0.023342382,0.0005143182,0.0027254557,0.0007411393,0.00060882897,0.013490431],"category_scores_gemma":[0.013434447,0.00051701884,0.0013185996,0.022589287,0.00031039648,0.0039447644,0.0015749998,0.00095511565,0.0043487176],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010943463,0.0003122342,0.049723577,0.005591067,0.00089651236,0.00065647595,0.0042673363,0.011529311,0.014060658,0.015999606,0.22258714,0.6732817],"study_design_scores_gemma":[0.00033376296,0.0008752237,0.124507956,0.0009649586,0.000654594,0.0019396069,0.0025559019,0.15110202,0.03155489,0.03733095,0.64775527,0.00042483513],"about_ca_topic_score_codex":0.006905898,"about_ca_topic_score_gemma":0.009472819,"teacher_disagreement_score":0.9766576,"about_ca_system_score_codex":0.0008979996,"about_ca_system_score_gemma":0.001285879,"threshold_uncertainty_score":0.045130014},"labels":[],"label_agreement":null},{"id":"W4366121787","doi":"10.2196/44876","title":"Identifying Patient Populations in Texts Describing Drug Approvals Through Deep Learning–Based Information Extraction: Development of a Natural Language Processing Algorithm","year":2023,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"AstraZeneca","keywords":"Computer science; Artificial intelligence; Machine learning; Deep learning; Information extraction; Population; Subject-matter expert; Automation; Natural language processing; Data science; Data mining; Expert system; Medicine; Engineering","score_opus":0.10278072382809395,"score_gpt":0.42817123600847745,"score_spread":0.3253905121803835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366121787","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02858363,0.0004933216,0.942929,0.0009516975,0.00010099569,0.0009135389,0.0074154367,0.017379662,0.001232625],"genre_scores_gemma":[0.07570115,0.000201672,0.90963495,0.00026191445,0.00005736787,0.0005771567,0.0121617345,0.00022890908,0.001175141],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9979298,0.00042675668,0.00050261954,0.0006851034,0.00034477745,0.0001109775],"domain_scores_gemma":[0.99364257,0.0044547804,0.0005259156,0.00042294932,0.0008269632,0.00012684976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025643262,0.0015731623,0.00086066936,0.0052902317,0.0006577875,0.0018082837,0.0016021294,0.0015987592,0.0035497013],"category_scores_gemma":[0.007953006,0.00063611387,0.0016915825,0.002199756,0.0006648464,0.0026884698,0.0015972407,0.0024755523,0.0032671667],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030328028,0.0004169837,0.007600544,0.00082389585,0.00014328185,0.0005404846,0.00064484193,0.029368065,0.022363339,0.0041183163,0.017849283,0.9158277],"study_design_scores_gemma":[0.00010624016,0.0001719646,0.0041717896,0.0002152606,0.000117391624,0.0007632036,0.00048825188,0.9186573,0.03597323,0.018355222,0.020909572,0.00007051627],"about_ca_topic_score_codex":0.004822013,"about_ca_topic_score_gemma":0.0068717203,"teacher_disagreement_score":0.0052902317,"about_ca_system_score_codex":0.0013407724,"about_ca_system_score_gemma":0.00265528,"threshold_uncertainty_score":0.013561606},"labels":[],"label_agreement":null},{"id":"W4366307907","doi":"10.2196/44835","title":"Natural Language Processing for Clinical Laboratory Data Repository Systems: Implementation and Evaluation for Respiratory Viruses","year":2023,"lang":"en","type":"article","venue":"JMIR AI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Sunnybrook Hospital; Sinai Health System; Vector Institute; Public Health Ontario; York University; University Health Network; University of Toronto","funders":"Canadian Institutes of Health Research; Hospital for Sick Children","keywords":"Computer science; Artificial intelligence; Natural language processing; Generalizability theory; Parsing; Classifier (UML); Machine learning; Information extraction; Task (project management)","score_opus":0.1563436978837656,"score_gpt":0.5307190196558628,"score_spread":0.3743753217720972,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4366307907","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7272162,0.0017781182,0.17627814,0.0020564888,0.00040037508,0.0047657625,0.005966243,0.07561905,0.005919729],"genre_scores_gemma":[0.6406522,0.0008408894,0.34147027,0.00058683864,0.00006123612,0.0017706483,0.010968541,0.0008535244,0.0027958401],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.996345,0.0012883132,0.000448976,0.00090819417,0.0008010579,0.00020857622],"domain_scores_gemma":[0.98994815,0.0064112023,0.00043971807,0.0008638227,0.0018690601,0.00046806358],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007174647,0.0009995921,0.00060461776,0.0010818596,0.00067974895,0.0014313578,0.0035055636,0.0014608268,0.0033778625],"category_scores_gemma":[0.018562445,0.0005005164,0.00071076525,0.0010493202,0.00077755126,0.0027499998,0.001896624,0.0018534457,0.0016871273],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004484861,0.005051716,0.026202483,0.0035657983,0.00069832517,0.0018038785,0.0027572438,0.11880503,0.038437672,0.0023635675,0.04794212,0.7478873],"study_design_scores_gemma":[0.00085423706,0.0011771715,0.009609471,0.0001646344,0.00019142817,0.00044375833,0.00080525974,0.93651235,0.0351666,0.0022501845,0.012712404,0.00011248182],"about_ca_topic_score_codex":0.018411856,"about_ca_topic_score_gemma":0.014442605,"teacher_disagreement_score":0.018411856,"about_ca_system_score_codex":0.0021035993,"about_ca_system_score_gemma":0.0030359006,"threshold_uncertainty_score":0.0379436},"labels":[],"label_agreement":null},{"id":"W4367845327","doi":"10.31234/osf.io/rxw9b","title":"Quantifying informativeness of names in visual space","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Agencia Estatal de Investigación; Natural Sciences and Engineering Research Council of Canada; Ministerio de Ciencia e Innovación; Australian Research Council; European Commission","keywords":"Lexicon; Computer science; Natural language processing; Optimal distinctiveness theory; Artificial intelligence; Set (abstract data type); Entropy (arrow of time); Semantic similarity; Space (punctuation); Measure (data warehouse); Object (grammar); Information retrieval; Psychology; Data mining","score_opus":0.07582923238812853,"score_gpt":0.3740497067090218,"score_spread":0.29822047432089327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4367845327","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5943253,0.0008386418,0.39885977,0.00045615298,0.000022485528,0.000032578264,0.00048106405,0.00038743194,0.0045966073],"genre_scores_gemma":[0.98169875,0.00018946154,0.01737252,0.000060733044,0.000040825675,0.000023506407,0.00028382504,0.00005141366,0.00027898807],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9979195,0.0007176958,0.00016225596,0.0005455942,0.00051676,0.0001381721],"domain_scores_gemma":[0.98008174,0.014505755,0.002634284,0.0015175986,0.00093233597,0.00032821123],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028011042,0.000416877,0.0006920175,0.0047125644,0.00058751943,0.002918428,0.0006630411,0.001191915,0.0009632118],"category_scores_gemma":[0.029474033,0.00046438898,0.00059374917,0.0031711026,0.0026997405,0.0061546043,0.0021029776,0.000893595,0.00020261935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018709458,0.0002182529,0.12576152,0.0009735169,0.0007034363,0.0011217402,0.006659078,0.16523562,0.1334872,0.24733382,0.0023505744,0.31428424],"study_design_scores_gemma":[0.000032744723,0.00021838406,0.08664522,0.000105432955,0.00023761498,0.0014957055,0.00139663,0.39938477,0.025747845,0.4817159,0.0028496902,0.00017011612],"about_ca_topic_score_codex":0.0009000984,"about_ca_topic_score_gemma":0.0007725386,"teacher_disagreement_score":0.0047125644,"about_ca_system_score_codex":0.0008765207,"about_ca_system_score_gemma":0.00028629834,"threshold_uncertainty_score":0.01481384},"labels":[],"label_agreement":null},{"id":"W4372352781","doi":"10.1016/s0992-5945(23)00080-6","title":"Numérique en santé : les biologistes au cœur des échanges de données","year":2023,"lang":"fr","type":"article","venue":"Option/Bio","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"World Federation of Science Journalists","funders":"","keywords":"Political science; Philosophy","score_opus":0.045293253581748655,"score_gpt":0.3099181838178637,"score_spread":0.264624930236115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4372352781","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039311677,0.020350827,0.75883526,0.08484998,0.0019806593,0.0003316615,0.007461844,0.004895543,0.08198254],"genre_scores_gemma":[0.28325087,0.012434684,0.6684799,0.004341029,0.00090551854,0.00040728928,0.005635266,0.00079946226,0.02374599],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9867013,0.004960899,0.0012277381,0.0016526221,0.0050856844,0.00037163953],"domain_scores_gemma":[0.98203456,0.012178654,0.0011974049,0.0014104088,0.0024671524,0.00071184686],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014027363,0.00082425674,0.0011730554,0.006869166,0.0019132929,0.013144205,0.0017288622,0.0028646118,0.005533593],"category_scores_gemma":[0.041467372,0.00066220085,0.0014442042,0.008057203,0.0043362426,0.017938761,0.004282639,0.0047624446,0.0018834466],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020972567,0.000108131026,0.008623195,0.0014676186,0.00017933841,0.00050334155,0.0037246158,0.0074968818,0.0036588646,0.662262,0.023866288,0.28789997],"study_design_scores_gemma":[0.00006713797,0.000073435716,0.008610301,0.0017718966,0.00013042358,0.001402463,0.004801327,0.061478008,0.0059200227,0.5088685,0.406712,0.0001644991],"about_ca_topic_score_codex":0.025204591,"about_ca_topic_score_gemma":0.014354879,"teacher_disagreement_score":0.025204591,"about_ca_system_score_codex":0.006138066,"about_ca_system_score_gemma":0.011212977,"threshold_uncertainty_score":0.074184656},"labels":[],"label_agreement":null},{"id":"W4375858571","doi":"10.32473/flairs.36.133256","title":"Identifying Protein-Protein Interaction using Tree-Transformers and Heterogeneous Graph Neural Network","year":2023,"lang":"en","type":"article","venue":"Proceedings of the ... International Florida Artificial Intelligence Research Society Conference","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Computer science; Artificial neural network; Transformer; Graph; Artificial intelligence; Theoretical computer science; Engineering; Electrical engineering","score_opus":0.17353665750856861,"score_gpt":0.3955412738069915,"score_spread":0.22200461629842289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4375858571","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28890958,0.0014216139,0.699088,0.0005661432,0.0001059385,0.0001444211,0.0010279249,0.0030572596,0.0056791715],"genre_scores_gemma":[0.9300618,0.00036848796,0.065344095,0.0001377027,0.000043499836,0.000056972338,0.0014212461,0.000065959415,0.0025002602],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99972767,0.00006625918,0.000012106772,0.000107410895,0.00004456805,0.000041926025],"domain_scores_gemma":[0.9994211,0.00033115392,0.00006651099,0.00005701805,0.000094716714,0.000029543475],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004949659,0.00067691616,0.00057640346,0.0019244858,0.0003494544,0.0006431061,0.0010622259,0.0008339798,0.0013044756],"category_scores_gemma":[0.0014916011,0.00026001284,0.00079845596,0.0015739219,0.0003895489,0.0017311886,0.00064114865,0.0006862718,0.0004787938],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052127946,0.0004547925,0.00805926,0.00020648405,0.00026915167,0.0005462847,0.00016360136,0.5660139,0.021466494,0.01587787,0.006911035,0.3795099],"study_design_scores_gemma":[0.0000028364834,0.000012256887,0.00043457528,0.0000016476739,0.00000978928,0.000014789395,0.000007016502,0.99489474,0.0006367714,0.0038307472,0.0001520067,0.0000027401861],"about_ca_topic_score_codex":0.012119324,"about_ca_topic_score_gemma":0.019246636,"teacher_disagreement_score":0.012119324,"about_ca_system_score_codex":0.0010494736,"about_ca_system_score_gemma":0.0005308927,"threshold_uncertainty_score":0.024097562},"labels":[],"label_agreement":null},{"id":"W4376104255","doi":"10.1101/2023.05.09.539725","title":"NGBO: Introducing -omics metadata to biobanking ontology","year":2023,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Genome British Columbia; Simon Fraser University; University of British Columbia","funders":"","keywords":"Biobank; Ontology; Computer science; Metadata; Data science; Discoverability; Open Biomedical Ontologies; Ontology-based data integration; Data integration; Information retrieval; Data management; World Wide Web; Semantic Web; Data mining; Suggested Upper Merged Ontology; Bioinformatics","score_opus":0.027166110540779533,"score_gpt":0.26091256741323826,"score_spread":0.23374645687245874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376104255","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042308164,0.0004928333,0.9557475,0.005434804,0.00061523705,0.0009635488,0.006301331,0.009193573,0.017020417],"genre_scores_gemma":[0.027655749,0.00089477136,0.9468501,0.0021752869,0.00018105285,0.0008516279,0.015080675,0.001491538,0.0048193075],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9941736,0.0017282349,0.0010840942,0.0006945584,0.0019166493,0.00040294847],"domain_scores_gemma":[0.99266934,0.0023602166,0.0005457255,0.0018192567,0.0020595917,0.00054576475],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012108716,0.0007642945,0.00059282803,0.0052044108,0.001970647,0.005479038,0.0026060566,0.0018714827,0.003823973],"category_scores_gemma":[0.012436334,0.000782426,0.002019926,0.0038796067,0.0021695204,0.010010782,0.0070872824,0.004054057,0.002259169],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014693967,0.00028779352,0.004747285,0.0015936581,0.00013007721,0.0013776495,0.004253493,0.0104895495,0.010995607,0.69985014,0.086980954,0.1791468],"study_design_scores_gemma":[0.000033135046,0.00002173949,0.0010458523,0.0008813013,0.00005499716,0.0003886704,0.00074245426,0.028355198,0.0049603344,0.11612724,0.84730977,0.00007936572],"about_ca_topic_score_codex":0.022304673,"about_ca_topic_score_gemma":0.022838673,"teacher_disagreement_score":0.022304673,"about_ca_system_score_codex":0.0046692626,"about_ca_system_score_gemma":0.009244273,"threshold_uncertainty_score":0.06403774},"labels":[],"label_agreement":null},{"id":"W4376130995","doi":"10.1016/j.jbi.2023.104384","title":"Deep learning to refine the identification of high-quality clinical research articles from the biomedical literature: Performance evaluation","year":2023,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hamilton Health Sciences; McMaster University; Impact","funders":"","keywords":"Computer science; Identification (biology); Quality (philosophy); Data science; Artificial intelligence; Deep learning","score_opus":0.12669730255498995,"score_gpt":0.452121140548754,"score_spread":0.32542383799376406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376130995","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7595867,0.023218004,0.18710993,0.004876989,0.0005854566,0.0013868969,0.008602049,0.0074705402,0.0071632927],"genre_scores_gemma":[0.88993686,0.0018178072,0.09773074,0.00064047094,0.00012165675,0.00032346058,0.007739665,0.00009250163,0.0015968733],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9962708,0.0017586245,0.00060620403,0.0005148097,0.0006370202,0.00021249414],"domain_scores_gemma":[0.9762408,0.01709487,0.0014172608,0.0016143135,0.0030492542,0.0005834652],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.018101895,0.0016594424,0.0014825714,0.003824393,0.00046204784,0.0016664164,0.0015701626,0.001812255,0.0016062411],"category_scores_gemma":[0.03639668,0.00054457935,0.0013888169,0.0018842205,0.00052341336,0.0017452786,0.0022953684,0.0022468048,0.0006936986],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0046616774,0.0013849422,0.06679279,0.0031263656,0.0026942913,0.0002066777,0.00026708943,0.2558163,0.0061913794,0.0014148743,0.011951359,0.6454923],"study_design_scores_gemma":[0.00036291592,0.00088111067,0.007640282,0.0003327513,0.0005519166,0.00015773892,0.000101049794,0.9768395,0.008576254,0.002464894,0.0020471115,0.00004459136],"about_ca_topic_score_codex":0.010888202,"about_ca_topic_score_gemma":0.012387127,"teacher_disagreement_score":0.9818981,"about_ca_system_score_codex":0.0019601767,"about_ca_system_score_gemma":0.003594494,"threshold_uncertainty_score":0.095733166},"labels":[],"label_agreement":null},{"id":"W4376638894","doi":"10.18280/isi.280206","title":"Metadata Analysis to Get Insight into Drug Resistant Ovarian Cancer","year":2023,"lang":"fr","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Science and Engineering Research Board","keywords":"Metadata; Ovarian cancer; Drug; Cancer; Computer science; World Wide Web; Oncology; Information retrieval; Medicine; Internal medicine; Pharmacology","score_opus":0.017514098815168674,"score_gpt":0.27572453438547934,"score_spread":0.25821043557031065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376638894","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25433388,0.013083412,0.29752436,0.0041680583,0.0005510955,0.0011433219,0.38588628,0.018214704,0.025094824],"genre_scores_gemma":[0.43301892,0.005653995,0.22986642,0.0006125089,0.00020245569,0.0005718016,0.3244179,0.00058843597,0.0050675482],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99879456,0.000148354,0.00029171142,0.0002806975,0.00037830882,0.000106329215],"domain_scores_gemma":[0.9972326,0.0008871329,0.000571687,0.0004544128,0.00069432793,0.00015984007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001257695,0.0005569134,0.00054946746,0.012434133,0.0006349234,0.0018381062,0.0005997745,0.0006216772,0.002604119],"category_scores_gemma":[0.0045957523,0.00014417035,0.00087425276,0.008601861,0.00034212266,0.0018536063,0.0011091141,0.00048720045,0.0013785851],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012496863,0.00043822956,0.1805656,0.00490942,0.00061711343,0.0051162792,0.0023715976,0.010101754,0.08519505,0.027834535,0.057418045,0.62418276],"study_design_scores_gemma":[0.00010447993,0.00049272197,0.19869244,0.0016484763,0.001068976,0.0064489273,0.0068174684,0.10524804,0.06752196,0.06571007,0.5459231,0.0003233181],"about_ca_topic_score_codex":0.0070459484,"about_ca_topic_score_gemma":0.008455892,"teacher_disagreement_score":0.012434133,"about_ca_system_score_codex":0.00097247906,"about_ca_system_score_gemma":0.0017085196,"threshold_uncertainty_score":0.014009893},"labels":[],"label_agreement":null},{"id":"W4376642717","doi":"10.1093/jamiaopen/ooad032","title":"A metadata framework for computational phenotypes","year":2023,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"National Institute of General Medical Sciences; National Center for Advancing Translational Sciences; National Human Genome Research Institute; Cincinnati Children's Hospital Medical Center","keywords":"Metadata; Computer science; Phenotype; Information retrieval; Computational biology; World Wide Web; Biology; Genetics","score_opus":0.06756581596636699,"score_gpt":0.3844595889207631,"score_spread":0.31689377295439614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4376642717","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014949234,0.0004020561,0.96740973,0.0036290854,0.00011708112,0.0013533756,0.002668715,0.0026854465,0.006785299],"genre_scores_gemma":[0.0667296,0.00021252221,0.9263961,0.0003690982,0.000049426202,0.001075007,0.0039957813,0.00027059097,0.00090190925],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9742656,0.014261747,0.0049086064,0.0023696246,0.003737554,0.00045691934],"domain_scores_gemma":[0.9050062,0.05527462,0.008080177,0.013872398,0.015630547,0.0021361103],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04911267,0.001023486,0.0009606077,0.013288254,0.0037262198,0.008420734,0.0035918546,0.0017956985,0.0032199123],"category_scores_gemma":[0.08937038,0.00087160064,0.0021029117,0.009345523,0.004286005,0.020218572,0.008251883,0.0025290824,0.0010698666],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039147396,0.00055066735,0.01919254,0.0016298238,0.00014990653,0.0004930379,0.026746964,0.008832496,0.0064001395,0.5846563,0.01298979,0.33796674],"study_design_scores_gemma":[0.00017498109,0.00045469395,0.01130217,0.0028175015,0.00030387114,0.0010948112,0.023613904,0.07600491,0.008638986,0.5949529,0.2802647,0.0003767051],"about_ca_topic_score_codex":0.010014601,"about_ca_topic_score_gemma":0.01303169,"teacher_disagreement_score":0.04911267,"about_ca_system_score_codex":0.004296074,"about_ca_system_score_gemma":0.008807824,"threshold_uncertainty_score":0.2597357},"labels":[],"label_agreement":null},{"id":"W4378746428","doi":"10.1177/14604582231180226","title":"A survey of epidemic management data models","year":2023,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Ministry of Science and Technology of the People's Republic of China; Department of Science and Technology of Shandong Province","keywords":"Interoperability; Computer science; Data science; Semantic interoperability; Data management; Controlled vocabulary; Knowledge management; Linked data; Ontology; Information retrieval; Semantic Web; World Wide Web; Data mining","score_opus":0.23603838108377093,"score_gpt":0.4104553707176358,"score_spread":0.1744169896338649,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378746428","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028823284,0.16204639,0.6633913,0.048951715,0.0010821768,0.001292176,0.025387857,0.0034008378,0.06562435],"genre_scores_gemma":[0.18008989,0.2390717,0.5135431,0.0071209376,0.0009679557,0.0015018206,0.04814642,0.0010169641,0.008541165],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.9874264,0.0026842372,0.003086864,0.0012065482,0.0051628654,0.0004330856],"domain_scores_gemma":[0.9765564,0.015254677,0.0015011566,0.002450343,0.0038706015,0.00036688062],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015469636,0.0010472055,0.0014925196,0.010132755,0.0014685603,0.007286552,0.0040374566,0.0019691105,0.0029833948],"category_scores_gemma":[0.030707493,0.0011627825,0.003498847,0.017226655,0.0017700857,0.014376713,0.0034616713,0.0027865241,0.001094971],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014305003,0.00018169527,0.0111576645,0.0042739217,0.00027521857,0.00047520138,0.0013120756,0.040353548,0.0011090342,0.6008837,0.04336648,0.29646835],"study_design_scores_gemma":[0.000028791928,0.0000758934,0.00352663,0.003033189,0.00012514017,0.0010313143,0.001188397,0.07473613,0.0014045268,0.21752053,0.69719464,0.00013492042],"about_ca_topic_score_codex":0.017516876,"about_ca_topic_score_gemma":0.008955082,"teacher_disagreement_score":0.017516876,"about_ca_system_score_codex":0.0055887746,"about_ca_system_score_gemma":0.006678598,"threshold_uncertainty_score":0.08181226},"labels":[],"label_agreement":null},{"id":"W4378901016","doi":"10.2196/45496","title":"Interoperable, Domain-Specific Extensions for the German Corona Consensus (GECCO) COVID-19 Research Data Set Using an Interdisciplinary, Consensus-Based Workflow: Data Set Development Study","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Interoperability; Workflow; Computer science; Data science; Domain (mathematical analysis); USable; World Wide Web; Database","score_opus":0.3909576814881041,"score_gpt":0.5202761976731551,"score_spread":0.12931851618505102,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378901016","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3644832,0.0010718678,0.47198206,0.004319248,0.00031074602,0.033520434,0.106527984,0.0055528698,0.012231586],"genre_scores_gemma":[0.14909557,0.00021611391,0.7223132,0.0005212343,0.000035079,0.014630112,0.111923255,0.00043202136,0.00083340076],"study_design_codex":"design_other","study_design_gemma":"design_other","domain_scores_codex":[0.96156526,0.018456196,0.0079195,0.005641865,0.005552638,0.0008645012],"domain_scores_gemma":[0.8666635,0.07316812,0.008565403,0.024077566,0.024055291,0.0034701154],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.067548886,0.00078351685,0.0007118362,0.0064414707,0.0023848123,0.0037681246,0.0029095684,0.001709522,0.0014802796],"category_scores_gemma":[0.123238385,0.00081747194,0.0024143113,0.0061051724,0.0017626347,0.003800127,0.007476072,0.0023890997,0.00077339774],"study_design_candidate":"design_other","study_design_consensus":"design_other","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030029276,0.0038936425,0.2811427,0.009812569,0.0011162146,0.002974658,0.043148767,0.06655029,0.02350629,0.11365042,0.08928718,0.36191425],"study_design_scores_gemma":[0.0023280585,0.0015517214,0.17285597,0.0050671757,0.0010348726,0.0020841223,0.026053479,0.2019141,0.052783195,0.06861645,0.46479762,0.0009133146],"about_ca_topic_score_codex":0.014502646,"about_ca_topic_score_gemma":0.018761527,"teacher_disagreement_score":0.99709046,"about_ca_system_score_codex":0.005690956,"about_ca_system_score_gemma":0.016481241,"threshold_uncertainty_score":0.35723692},"labels":[],"label_agreement":null},{"id":"W4379184839","doi":"10.1371/journal.pone.0286728","title":"OntoTrek: 3D visualization of application ontology class hierarchies","year":2023,"lang":"en","type":"article","venue":"PLoS ONE","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Simon Fraser University","funders":"Natural Sciences and Engineering Research Council of Canada; Agricultural Research Service; Genome Canada; U.S. Department of Agriculture","keywords":"Visualization; Ontology; Class (philosophy); Computer science; Computational biology; Data science; World Wide Web; Evolutionary biology; Biology; Data mining; Artificial intelligence; Epistemology","score_opus":0.038584063951070106,"score_gpt":0.28252470161466575,"score_spread":0.24394063766359564,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379184839","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06765545,0.0007917775,0.7539889,0.0033225582,0.0006515097,0.00040807368,0.023166534,0.101363294,0.048651867],"genre_scores_gemma":[0.4169751,0.0016665047,0.52265483,0.0012218878,0.00015312356,0.0009166943,0.016275851,0.017563716,0.022572378],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9997396,0.0000570397,0.00002099306,0.000037738533,0.00010641581,0.00003826956],"domain_scores_gemma":[0.9984794,0.00085002603,0.000094940144,0.00016619472,0.00026833764,0.00014111976],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008578165,0.0009175805,0.00036624793,0.0019546156,0.00059077574,0.0023893623,0.00078768807,0.00091930135,0.019477231],"category_scores_gemma":[0.0029118708,0.0003664409,0.00065609545,0.0010901635,0.00044317008,0.0021892064,0.003081333,0.0013690491,0.0021458932],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016432415,0.00046556952,0.014045437,0.0029165968,0.00028547374,0.0030120437,0.024680307,0.024751103,0.08319691,0.08227494,0.3588009,0.40392742],"study_design_scores_gemma":[0.00018778052,0.00018873239,0.017261868,0.00075925444,0.000113285714,0.0014643795,0.0044381185,0.17698999,0.039155588,0.04777638,0.71136487,0.00029973546],"about_ca_topic_score_codex":0.0057904683,"about_ca_topic_score_gemma":0.009108385,"teacher_disagreement_score":0.019477231,"about_ca_system_score_codex":0.0005002214,"about_ca_system_score_gemma":0.0009774826,"threshold_uncertainty_score":0.06515777},"labels":[],"label_agreement":null},{"id":"W4379341789","doi":"10.1200/jco.2023.41.16_suppl.e18653","title":"Medical terminology and interpretation of results in plain language summaries published by oncology journals.","year":2023,"lang":"en","type":"article","venue":"Journal of Clinical Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Syncrude (Canada)","funders":"","keywords":"Readability; Terminology; Medicine; Medical terminology; Context (archaeology); Medical physics; Computer science; Linguistics","score_opus":0.05592858290550649,"score_gpt":0.4539058598675796,"score_spread":0.3979772769620731,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379341789","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5499451,0.07862759,0.1457137,0.032153063,0.014955565,0.021863509,0.09645788,0.00687728,0.053406317],"genre_scores_gemma":[0.7493711,0.014295117,0.19456254,0.003126436,0.0029624943,0.012500439,0.017384924,0.0010844014,0.0047124354],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.7347519,0.11776308,0.11549161,0.0053713256,0.025422134,0.0011999252],"domain_scores_gemma":[0.20154911,0.53126824,0.16287144,0.024660727,0.07721707,0.0024334923],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.16167527,0.0014552535,0.0016023818,0.028046321,0.001327086,0.006437706,0.0016109155,0.0013428311,0.016540196],"category_scores_gemma":[0.60597324,0.0009173673,0.0021999131,0.01714667,0.0023864384,0.004774538,0.0041635684,0.0009400285,0.0049537034],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009023108,0.00029089348,0.22048281,0.10044271,0.0024416482,0.0028744892,0.03286005,0.0016525359,0.0148651665,0.007316148,0.12869464,0.47905585],"study_design_scores_gemma":[0.001038353,0.0031182496,0.5256487,0.043402985,0.0034577516,0.01189364,0.02477193,0.0072717373,0.024344726,0.02113526,0.3328963,0.0010204007],"about_ca_topic_score_codex":0.00046305868,"about_ca_topic_score_gemma":0.00075824576,"teacher_disagreement_score":0.9935623,"about_ca_system_score_codex":0.0019923134,"about_ca_system_score_gemma":0.0047855307,"threshold_uncertainty_score":0.8550308},"labels":[],"label_agreement":null},{"id":"W4379652851","doi":"10.1136/annrheumdis-2023-eular.6231","title":"AB1767-HPR DOCUMENT SEARCH IN LARGE RHEUMATOLOGY DATABASES: ADVANCED KEYWORD QUERIES TO SELECT HOMOGENEOUS PHENOTYPES","year":2023,"lang":"en","type":"preprint","venue":"Annals of the Rheumatic Diseases","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; University of British Columbia","funders":"","keywords":"Computer science; Homogeneous; Keyword search; Information retrieval; Database; World Wide Web; Mathematics","score_opus":0.05282673366597263,"score_gpt":0.3598819778971527,"score_spread":0.3070552442311801,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4379652851","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.070051335,0.0030410686,0.17946923,0.002981494,0.00049572665,0.0007768263,0.47137645,0.24212164,0.029686194],"genre_scores_gemma":[0.120365396,0.0014365718,0.27478576,0.0005674687,0.0002592558,0.000481161,0.5802231,0.009906476,0.011974789],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987721,0.00024599128,0.00017882188,0.00029292895,0.00040141676,0.000108748274],"domain_scores_gemma":[0.99787223,0.001139321,0.000102874495,0.00044105717,0.00031947013,0.00012495021],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015617566,0.0014795688,0.0010396284,0.0044283904,0.0005405587,0.002370589,0.0011132674,0.0016663634,0.031193033],"category_scores_gemma":[0.007959223,0.0005664085,0.0010457474,0.0035815686,0.00031675925,0.0021616318,0.0015824422,0.0006501426,0.018659554],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0026666166,0.0005320487,0.006651891,0.0030381181,0.0006226449,0.001182949,0.00052676693,0.007660562,0.028407414,0.014018703,0.65019166,0.2845007],"study_design_scores_gemma":[0.0029736063,0.00049542764,0.020753754,0.00046015263,0.00059017417,0.0027529309,0.0010719694,0.2521627,0.058220748,0.048242223,0.6120584,0.00021790409],"about_ca_topic_score_codex":0.005299026,"about_ca_topic_score_gemma":0.0064653517,"teacher_disagreement_score":0.031193033,"about_ca_system_score_codex":0.0006121257,"about_ca_system_score_gemma":0.0014405005,"threshold_uncertainty_score":0.10435116},"labels":[],"label_agreement":null},{"id":"W4380989093","doi":"10.2196/47434","title":"A Deep Learning Model for the Normalization of Institution Names by Multisource Literature Feature Fusion: Algorithm Development Study","year":2023,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Chinese Academy of Medical Sciences","keywords":"Normalization (sociology); Computer science; Artificial intelligence; Institution; Deep learning; Natural language processing; Scopus; Machine learning; Information retrieval; Political science; Law","score_opus":0.03768942248781827,"score_gpt":0.3803296186251347,"score_spread":0.34264019613731644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4380989093","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04014363,0.0006274937,0.9540635,0.00048752932,0.0000827557,0.00013958375,0.00021006765,0.0023935668,0.0018518898],"genre_scores_gemma":[0.6101101,0.0006739929,0.37969372,0.00038908774,0.000075864344,0.0005178689,0.0013626708,0.00012803324,0.0070487224],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996569,0.00005411114,0.000026246666,0.00012158021,0.000080594,0.00006058976],"domain_scores_gemma":[0.9994142,0.00018812391,0.000045886554,0.00004858257,0.00026928986,0.000033971068],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012962506,0.00087299786,0.0009002181,0.0009620709,0.00046571647,0.000842475,0.0017051924,0.0011888229,0.0024478724],"category_scores_gemma":[0.0022708944,0.00046740624,0.00091019523,0.0011684273,0.00040726454,0.0013874585,0.0012533541,0.0017233323,0.00081709656],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014901215,0.00019610491,0.0039660134,0.00010601761,0.00014133747,0.00013232008,0.00009531476,0.42367926,0.0054627345,0.0049887644,0.0046931924,0.55639],"study_design_scores_gemma":[0.0000047681997,0.000014275049,0.0001556051,0.0000047487456,0.000011059247,0.000014350559,0.000007870045,0.9979882,0.0008122663,0.0007302617,0.00025317853,0.0000032835524],"about_ca_topic_score_codex":0.016324202,"about_ca_topic_score_gemma":0.012330727,"teacher_disagreement_score":0.016324202,"about_ca_system_score_codex":0.0013941408,"about_ca_system_score_gemma":0.0025134173,"threshold_uncertainty_score":0.032458365},"labels":[],"label_agreement":null},{"id":"W4381189786","doi":"10.1038/s41559-023-02104-x","title":"Change in biological nomenclature is overdue and possible","year":2023,"lang":"en","type":"letter","venue":"Nature Ecology & Evolution","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":21,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg; University of Manitoba","funders":"","keywords":"Nomenclature; Computational biology; Biology; Data science; Evolutionary biology; Zoology; Computer science; Taxonomy (biology)","score_opus":0.02027283007638825,"score_gpt":0.28804139801561274,"score_spread":0.2677685679392245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381189786","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00018538721,0.0007813886,0.00060407206,0.970974,0.025794009,0.0000068434924,0.000034293385,0.000056065306,0.0015639132],"genre_scores_gemma":[0.0015046275,0.0006753399,0.001846944,0.9541707,0.03440444,0.00002317887,0.00007791419,0.000057943944,0.0072389124],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9694539,0.0047490075,0.0036470585,0.0039297543,0.015619705,0.0026005348],"domain_scores_gemma":[0.855226,0.07452574,0.007757738,0.008643239,0.033349775,0.02049752],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.039543778,0.00094505027,0.0025445132,0.00256832,0.007799717,0.012443172,0.0044534677,0.060688324,0.011717861],"category_scores_gemma":[0.10214101,0.0010779261,0.0032405744,0.0031363769,0.014972768,0.015689343,0.006819926,0.0734751,0.013546421],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000043234293,0.00004021395,0.00035932043,0.00006954573,0.000030305011,0.0004309333,0.0002236079,0.0000710288,0.0003291312,0.012808116,0.965733,0.01986163],"study_design_scores_gemma":[0.000045732126,0.000036313624,0.0010533436,0.00025651354,0.000029017616,0.00074869074,0.0004588321,0.00045431766,0.00015143829,0.038023666,0.9586822,0.00005993611],"about_ca_topic_score_codex":0.009482321,"about_ca_topic_score_gemma":0.030628458,"teacher_disagreement_score":0.060688324,"about_ca_system_score_codex":0.00881235,"about_ca_system_score_gemma":0.01624841,"threshold_uncertainty_score":0.20912999},"labels":[],"label_agreement":null},{"id":"W4381377622","doi":"10.1101/2023.06.18.23291567","title":"Machine learning to increase the efficiency of a literature surveillance system: a performance evaluation","year":2023,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Impact","funders":"","keywords":"Computer science; Machine learning; Critical appraisal; Artificial intelligence; MEDLINE; Set (abstract data type); Sensitivity (control systems); Binary classification; Information retrieval; Medical physics; Medicine; Support vector machine; Alternative medicine; Pathology","score_opus":0.023546279314569988,"score_gpt":0.2870682422268218,"score_spread":0.26352196291225183,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381377622","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6057006,0.012904524,0.32992023,0.0055544926,0.00055727153,0.0046095815,0.005137544,0.029111747,0.006503959],"genre_scores_gemma":[0.67714036,0.00067560986,0.3174545,0.0005542074,0.00012814642,0.0008477191,0.002488047,0.00028125523,0.00043018613],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9423803,0.03651693,0.008636491,0.0054302397,0.0061854096,0.0008506642],"domain_scores_gemma":[0.63669926,0.31054062,0.01562753,0.01608887,0.01889476,0.002148992],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.089243695,0.0018544234,0.0028409678,0.010228187,0.0013862092,0.0045645493,0.0024066658,0.0026577534,0.001902437],"category_scores_gemma":[0.23593546,0.0007500413,0.0026387272,0.00611454,0.0011613558,0.004766684,0.002853256,0.0017119275,0.0012826266],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008924104,0.0019051483,0.18144415,0.0044997833,0.0036595175,0.0003834,0.0013691948,0.085915074,0.007442376,0.002509013,0.013149241,0.6887989],"study_design_scores_gemma":[0.0007704928,0.0014105182,0.016740384,0.0005216432,0.0008845251,0.00042027584,0.00026245066,0.96439064,0.008265295,0.0037594738,0.0024744766,0.000099839686],"about_ca_topic_score_codex":0.0056382907,"about_ca_topic_score_gemma":0.004337159,"teacher_disagreement_score":0.9107563,"about_ca_system_score_codex":0.0024945983,"about_ca_system_score_gemma":0.0054245917,"threshold_uncertainty_score":0.4719714},"labels":[],"label_agreement":null},{"id":"W4381431186","doi":"10.3233/sw-233207","title":"Reuse of the FoodOn ontology in a knowledge base of food composition data","year":2023,"lang":"en","type":"article","venue":"Semantic Web","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"SPARQL; Leverage (statistics); Computer science; Ontology; Identifier; Knowledge base; Linked data; Food composition data; RDF; Reuse; Composition (language); World Wide Web; Information retrieval; Domain (mathematical analysis); Data science; Semantic Web; Food science; Artificial intelligence; Biology; Ecology","score_opus":0.05396335660828126,"score_gpt":0.3079433826309874,"score_spread":0.25398002602270614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4381431186","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027493432,0.00082956883,0.9116061,0.0017627841,0.0003711974,0.00089120097,0.015290968,0.016393555,0.02536125],"genre_scores_gemma":[0.080087915,0.0010124954,0.8683726,0.0010044599,0.00006419816,0.00041031823,0.04088636,0.0026020014,0.0055596246],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9956981,0.00064341276,0.00073496095,0.0011437996,0.001535003,0.0002447678],"domain_scores_gemma":[0.9841982,0.005029247,0.00092154247,0.0060684145,0.0032467388,0.00053579983],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073965373,0.0010678942,0.0011977941,0.0077548036,0.0017212583,0.0063358378,0.0031941086,0.0012218382,0.004020076],"category_scores_gemma":[0.01560957,0.0010935221,0.0026205808,0.0064637717,0.0015115574,0.011651078,0.005806561,0.0025414494,0.0016869861],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047523502,0.0015262789,0.022852404,0.003940936,0.0009063176,0.0039211083,0.0072119837,0.020158002,0.026265118,0.1510345,0.05955334,0.70215476],"study_design_scores_gemma":[0.000071594935,0.00010777511,0.008287119,0.0019599285,0.00062481716,0.0018834289,0.0021726629,0.07603825,0.039956134,0.08875909,0.77984715,0.00029205033],"about_ca_topic_score_codex":0.01917035,"about_ca_topic_score_gemma":0.03263307,"teacher_disagreement_score":0.01917035,"about_ca_system_score_codex":0.0029088682,"about_ca_system_score_gemma":0.005910338,"threshold_uncertainty_score":0.039117098},"labels":[],"label_agreement":null},{"id":"W4382362705","doi":"10.21203/rs.3.rs-3085472/v1","title":"Drug and natural health product data collection and curation in the Canadian Longitudinal Study on Aging (CLSA)","year":2023,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University; Dalhousie University; University of Toronto; McMaster University; Université de Sherbrooke; Western University","funders":"","keywords":"Data collection; Sample (material); Gold standard (test); Product (mathematics); Computer science; Drug; Medicine; Statistics; Mathematics; Pharmacology; Chemistry","score_opus":0.28500809479869177,"score_gpt":0.4962583178454439,"score_spread":0.21125022304675212,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4382362705","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.37741622,0.020729514,0.04475048,0.017372664,0.00080593134,0.01632148,0.46176928,0.0015986161,0.059235785],"genre_scores_gemma":[0.711411,0.011323783,0.11683203,0.0062405244,0.00032495943,0.008792563,0.13675323,0.00040791085,0.007914013],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9615908,0.010029953,0.0039478433,0.0037001523,0.01895574,0.0017753548],"domain_scores_gemma":[0.8934145,0.013453859,0.010152572,0.009993893,0.06940995,0.0035753464],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.039252978,0.0009788731,0.0008727941,0.008667808,0.0036669162,0.002694392,0.0038761846,0.0006554124,0.0034874058],"category_scores_gemma":[0.07490653,0.000584751,0.0009955472,0.012010999,0.0017376248,0.0009251676,0.0030163915,0.0011403598,0.00070218515],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050308043,0.00018841981,0.78494084,0.004072112,0.000633929,0.00021928023,0.003360441,0.001365359,0.0011583078,0.0042665787,0.08624296,0.113048695],"study_design_scores_gemma":[0.00012095532,0.00016223626,0.8891868,0.0024335673,0.00033957616,0.00014412773,0.0017959708,0.0025474518,0.0021888094,0.0010383086,0.09991969,0.00012249342],"about_ca_topic_score_codex":0.9700777,"about_ca_topic_score_gemma":0.96761954,"teacher_disagreement_score":0.9684277,"about_ca_system_score_codex":0.03157227,"about_ca_system_score_gemma":0.15205775,"threshold_uncertainty_score":0.22907388},"labels":[],"label_agreement":null},{"id":"W4383040509","doi":"10.3389/frai.2023.1137961","title":"Ontological how and why: action and objective of planned processes in the food domain","year":2023,"lang":"en","type":"review","venue":"Frontiers in Artificial Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Agricultural Research Service; Institut National de Recherche pour l'Agriculture, l'Alimentation et l'Environnement; Genome Canada; U.S. Department of Agriculture","keywords":"Computer science; Process (computing); Ontology; Field (mathematics); Automation; Domain (mathematical analysis); Flexibility (engineering); Interdependence; Data science; Artificial intelligence; Engineering","score_opus":0.12964164253569616,"score_gpt":0.36224472023481286,"score_spread":0.2326030776991167,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383040509","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015790601,0.9123692,0.03531782,0.0088402815,0.0009725339,0.000058533176,0.00031159705,0.00010466549,0.040446315],"genre_scores_gemma":[0.03725782,0.9268969,0.024145698,0.0034536386,0.00052613637,0.00013781,0.00074634707,0.000053900818,0.006781794],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994529,0.00019054915,0.00006083218,0.00008671594,0.0001719678,0.00003709076],"domain_scores_gemma":[0.9992563,0.00045325753,0.00008108159,0.00007140108,0.00010781375,0.000030079316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016600563,0.00081799785,0.00080245093,0.0023426274,0.00054168416,0.0027312941,0.0014688399,0.0017558376,0.0018819352],"category_scores_gemma":[0.0024167865,0.00030133486,0.00070764945,0.003291274,0.0031162347,0.0058172396,0.0014435779,0.002456906,0.00096395274],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038766204,0.000041767104,0.00042868007,0.006202073,0.000065674176,0.00018610076,0.0011207369,0.0009880043,0.0008650314,0.47599655,0.017284108,0.49678257],"study_design_scores_gemma":[0.0000068476193,0.000011645472,0.0011749733,0.004368956,0.00007165292,0.0007432099,0.0005895104,0.0007876167,0.0006978759,0.11858407,0.8729301,0.000033478063],"about_ca_topic_score_codex":0.006082839,"about_ca_topic_score_gemma":0.006754587,"teacher_disagreement_score":0.006082839,"about_ca_system_score_codex":0.0029041062,"about_ca_system_score_gemma":0.0034898045,"threshold_uncertainty_score":0.021070838},"labels":[],"label_agreement":null},{"id":"W4383098061","doi":"10.2196/48645","title":"Comprehensive Ontology of Fibroproliferative Diseases: Protocol for a Semantic Technology Study","year":2023,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Deutsche Forschungsgemeinschaft","keywords":"Ontology; Computer science; Protocol (science); Natural language processing; Data science; Medicine; Information retrieval; Pathology; Alternative medicine","score_opus":0.3056264003572566,"score_gpt":0.5911879511506735,"score_spread":0.2855615507934169,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383098061","genre_codex":"protocol","genre_gemma":"protocol","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":"protocol","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0057222717,0.0020435774,0.17478047,0.009180028,0.0015664243,0.7480402,0.029482206,0.0009820191,0.028202847],"genre_scores_gemma":[0.0065001636,0.0014354124,0.22659568,0.002097829,0.00015380567,0.74907243,0.010255736,0.00019638265,0.0036925613],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9470689,0.030081693,0.013535939,0.0027760663,0.0051512946,0.0013861281],"domain_scores_gemma":[0.88352156,0.060578197,0.0060923337,0.020302465,0.02698591,0.002519542],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11186147,0.001546784,0.0020458836,0.008271934,0.004183116,0.006265679,0.0034652147,0.0047493973,0.045897026],"category_scores_gemma":[0.12509622,0.0017410765,0.004501065,0.0077634407,0.0044296836,0.007329038,0.009804126,0.0067218877,0.012091201],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004203493,0.0017155166,0.0037543958,0.07265023,0.0006562248,0.0029665616,0.04409833,0.0037053565,0.014563305,0.22768885,0.28304148,0.34095624],"study_design_scores_gemma":[0.0014626838,0.000330816,0.002402185,0.019625442,0.0004314306,0.00071982393,0.0072566574,0.0011934127,0.0040065767,0.04261796,0.9197357,0.00021725621],"about_ca_topic_score_codex":0.003306592,"about_ca_topic_score_gemma":0.0053642425,"teacher_disagreement_score":0.11186147,"about_ca_system_score_codex":0.008278798,"about_ca_system_score_gemma":0.0461884,"threshold_uncertainty_score":0.59158707},"labels":[],"label_agreement":null},{"id":"W4383618376","doi":"10.1093/jamiaopen/ooad046","title":"AnnoDash, a clinical terminology annotation dashboard","year":2023,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute for Clinical Evaluative Sciences; SickKids Foundation; Hospital for Sick Children","funders":"","keywords":"Computer science; SNOMED CT; Information retrieval; Terminology; Ontology; Annotation; Dashboard; Ranking (information retrieval); Interoperability; Interface (matter); Data science; World Wide Web; Artificial intelligence","score_opus":0.09907901780984156,"score_gpt":0.43523982134038697,"score_spread":0.3361608035305454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383618376","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012786541,0.0012654512,0.2111298,0.005131392,0.0024131725,0.0016446383,0.10140981,0.63457185,0.029647386],"genre_scores_gemma":[0.074332625,0.0021436631,0.53131723,0.0055615883,0.0011320683,0.0037133503,0.2550169,0.07056973,0.056212895],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9942978,0.0011368265,0.00089372165,0.0010767784,0.0023604434,0.00023444956],"domain_scores_gemma":[0.9696144,0.014766801,0.0023414101,0.00582096,0.0053795483,0.0020768985],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008827553,0.002771909,0.0013677947,0.00520833,0.001039726,0.0053721378,0.0034905872,0.0017867031,0.060391817],"category_scores_gemma":[0.03969884,0.0011517906,0.001282871,0.0038963899,0.00080661266,0.0063349493,0.007885338,0.0024719199,0.02614248],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018180985,0.00030762242,0.0047295005,0.0015089032,0.00014375233,0.0009231221,0.0012472797,0.0024024134,0.005256208,0.0073162857,0.6532911,0.32105565],"study_design_scores_gemma":[0.00040450186,0.00019345264,0.0054012565,0.00061824184,0.000079655205,0.000558254,0.00044650637,0.017394448,0.011285618,0.01356381,0.9498198,0.00023454987],"about_ca_topic_score_codex":0.0039310884,"about_ca_topic_score_gemma":0.0045720735,"teacher_disagreement_score":0.060391817,"about_ca_system_score_codex":0.0016179142,"about_ca_system_score_gemma":0.002561329,"threshold_uncertainty_score":0.20203072},"labels":[],"label_agreement":null},{"id":"W4383621006","doi":"10.1007/s00335-023-10005-4","title":"Bridging mouse and human anatomies; a knowledge-based approach to comparative anatomy for disease model phenotyping","year":2023,"lang":"en","type":"review","venue":"Mammalian Genome","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children; University of Toronto","funders":"Education, Audiovisual and Culture Executive Agency; Universitat Autònoma de Barcelona; European Commission","keywords":"Comparative anatomy; Biology; Human anatomy; Human disease; Anatomy; Human genetics; Genetics; Gene","score_opus":0.12727298173116183,"score_gpt":0.39116425498642193,"score_spread":0.2638912732552601,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4383621006","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00069178006,0.9718159,0.017269699,0.0019922503,0.00044796258,0.00008889504,0.0009611439,0.00025465954,0.0064777923],"genre_scores_gemma":[0.0045143776,0.96260965,0.027289169,0.0012441011,0.000259691,0.0001682767,0.001885701,0.00007805543,0.0019510278],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99881315,0.00031465827,0.0001921121,0.00016066084,0.000472696,0.000046739056],"domain_scores_gemma":[0.9970643,0.0020385801,0.00028514094,0.00014836305,0.00037007112,0.00009354873],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036595317,0.00084823515,0.0015186622,0.009745252,0.00036581347,0.0025601366,0.0018520736,0.0012936969,0.0036180601],"category_scores_gemma":[0.005105047,0.00039300064,0.0011210339,0.0072852215,0.0014211073,0.0034998583,0.0015774041,0.0021005964,0.0020904269],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000438363,0.000048723374,0.00035783177,0.022137793,0.00022060625,0.0002538119,0.00034765215,0.00040365872,0.0024115855,0.024309032,0.027024677,0.92244077],"study_design_scores_gemma":[0.000011839924,0.000039413568,0.001998565,0.011305782,0.00033474644,0.0013942275,0.00027877875,0.0003649098,0.0012156419,0.018711887,0.9643034,0.000040856798],"about_ca_topic_score_codex":0.0028296812,"about_ca_topic_score_gemma":0.003771771,"teacher_disagreement_score":0.009745252,"about_ca_system_score_codex":0.0013573613,"about_ca_system_score_gemma":0.002745189,"threshold_uncertainty_score":0.019353688},"labels":[],"label_agreement":null},{"id":"W4384154318","doi":"10.1101/2023.07.13.23292612","title":"The Medical Action Ontology: A Tool for Annotating and Analyzing Treatments and Clinical Management of Human Disease","year":2023,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Children's Hospital of Eastern Ontario","funders":"National Human Genome Research Institute; Medical Research Council; National Institutes of Health","keywords":"Ontology; Computer science; Open Biomedical Ontologies; Action (physics); Task (project management); Data science; Process ontology; Semantic Web; World Wide Web; Ontology alignment","score_opus":0.08964938169720421,"score_gpt":0.42367284596328686,"score_spread":0.33402346426608265,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384154318","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00378831,0.004094119,0.7020467,0.006387253,0.00086077646,0.0017573072,0.17480986,0.06040456,0.045851085],"genre_scores_gemma":[0.021891033,0.0054633874,0.79081553,0.0024335366,0.00030526254,0.0017313483,0.16470508,0.0059200255,0.006734793],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99653244,0.00087576575,0.00092419545,0.0005567367,0.0009473346,0.00016353502],"domain_scores_gemma":[0.99147147,0.0048992326,0.0009085605,0.0012985023,0.00090580934,0.00051648286],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005953809,0.0019611993,0.0011035643,0.014084542,0.0020228357,0.004436989,0.0027065908,0.0028318388,0.020513006],"category_scores_gemma":[0.016193252,0.0013425517,0.0032125912,0.008033514,0.0015270822,0.007495873,0.0072010146,0.0028492066,0.008725582],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003903962,0.00023851925,0.0051124226,0.0095288,0.0004747412,0.0018344529,0.0035550504,0.010041638,0.007768725,0.24710424,0.45421383,0.2597373],"study_design_scores_gemma":[0.00007366369,0.00003004272,0.001408712,0.0012867217,0.000097241915,0.0006855203,0.0003867251,0.00871823,0.0018863417,0.06014664,0.92519224,0.00008793708],"about_ca_topic_score_codex":0.013564982,"about_ca_topic_score_gemma":0.018434212,"teacher_disagreement_score":0.020513006,"about_ca_system_score_codex":0.0032609534,"about_ca_system_score_gemma":0.00817453,"threshold_uncertainty_score":0.06862283},"labels":[],"label_agreement":null},{"id":"W4384695050","doi":"10.3389/fninf.2023.1174156","title":"NIDM-Terms: community-based terminology management for improved neuroimaging dataset descriptions and query","year":2023,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"National Institute of Biomedical Imaging and Bioengineering; National Institute of Mental Health; Canada First Research Excellence Fund; Health Canada; National Institutes of Health; Fondation Brain Canada; McGill University","keywords":"Computer science; Neuroimaging; Terminology; Information retrieval; Neuroinformatics; Metadata; Directory; Data science; Data management; Data mining; World Wide Web","score_opus":0.032515970837142125,"score_gpt":0.28392109487871137,"score_spread":0.25140512404156923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4384695050","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036869743,0.0005775769,0.7109989,0.0021977057,0.00054082635,0.0017938716,0.07234932,0.20125896,0.0065958058],"genre_scores_gemma":[0.022078134,0.00061625225,0.7747924,0.0012466945,0.00016503305,0.0027771129,0.16652766,0.029198319,0.0025984263],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9833243,0.0036628454,0.005317339,0.0024605584,0.004578948,0.00065592874],"domain_scores_gemma":[0.9572425,0.016007373,0.003235949,0.0143935,0.0071371496,0.0019835355],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03345533,0.0031130044,0.0023430667,0.016279813,0.0032377879,0.010594953,0.0062151514,0.0028303096,0.014796441],"category_scores_gemma":[0.08419571,0.0025687083,0.0048415554,0.013909618,0.002099646,0.020796182,0.020758865,0.00549569,0.012501502],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012550164,0.000445279,0.009094124,0.0037895269,0.00049567234,0.0011108586,0.006637941,0.008784426,0.014962498,0.16034077,0.528872,0.26421195],"study_design_scores_gemma":[0.00034853155,0.00014722365,0.003288933,0.0011482186,0.00016960585,0.0007432524,0.0020644867,0.077979006,0.015719952,0.19183245,0.7060771,0.0004813339],"about_ca_topic_score_codex":0.014294224,"about_ca_topic_score_gemma":0.017737603,"teacher_disagreement_score":0.9665447,"about_ca_system_score_codex":0.005476464,"about_ca_system_score_gemma":0.008684882,"threshold_uncertainty_score":0.17693079},"labels":[],"label_agreement":null},{"id":"W4385475778","doi":"10.1093/bib/bbad267","title":"ReProMSig: an integrative platform for development and application of reproducible multivariable models for cancer prognosis supporting guideline-based transparent reporting","year":2023,"lang":"en","type":"article","venue":"Briefings in Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Cancer Research","funders":"Beijing Municipal Science and Technology Commission; Beijing Municipal Health Commission; Baidu","keywords":"Checklist; Guideline; Computer science; Resource (disambiguation); Multivariable calculus; Clinical Practice; Medicine; Medical physics; Data mining; Data science; Pathology; Family medicine; Psychology","score_opus":0.0899427042031709,"score_gpt":0.37564082302961993,"score_spread":0.285698118826449,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385475778","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0056709102,0.0018205026,0.44722903,0.0051072817,0.0007709231,0.0033161307,0.08668443,0.43697032,0.0124304155],"genre_scores_gemma":[0.052032188,0.0032297717,0.61799824,0.004331547,0.000923558,0.010132289,0.23650786,0.06316767,0.011676844],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97139883,0.010962854,0.007187182,0.0036210143,0.0059075006,0.0009225633],"domain_scores_gemma":[0.8661716,0.066884406,0.014731502,0.029757984,0.019357197,0.003097272],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.051311426,0.0035578231,0.0020810517,0.011804957,0.0010059855,0.008191772,0.005343385,0.0029629795,0.02661746],"category_scores_gemma":[0.13915089,0.001969668,0.0032258297,0.004741278,0.0011513113,0.006689422,0.010352355,0.003399025,0.029520608],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020011896,0.00040384702,0.017653031,0.005716015,0.00075388426,0.001756726,0.0026320827,0.0068021663,0.008972064,0.021937931,0.5048977,0.42647326],"study_design_scores_gemma":[0.0007795405,0.0004686488,0.014042465,0.0049945954,0.00047279263,0.0012888081,0.0007632622,0.0422372,0.024284808,0.06355172,0.8462717,0.00084446714],"about_ca_topic_score_codex":0.0035813048,"about_ca_topic_score_gemma":0.0023777876,"teacher_disagreement_score":0.94868857,"about_ca_system_score_codex":0.0017762394,"about_ca_system_score_gemma":0.008257317,"threshold_uncertainty_score":0.27136397},"labels":[],"label_agreement":null},{"id":"W4385486326","doi":"10.1007/978-3-031-33390-3_13","title":"Feature Engineering","year":2023,"lang":"en","type":"book-chapter","venue":"Statisctics and computing/Statistics and computing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Feature engineering; Feature (linguistics); Computer science; Artificial intelligence; Logarithm; Data mining; Machine learning; Mathematics; Deep learning","score_opus":0.012205269216241557,"score_gpt":0.24356324849001937,"score_spread":0.23135797927377783,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385486326","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011133367,0.0003789352,0.8955999,0.00053930684,0.0003894677,0.00049626146,0.008065662,0.03637066,0.04702638],"genre_scores_gemma":[0.17686251,0.0008475803,0.61947656,0.0008812847,0.00021137022,0.0008848712,0.04176397,0.008518859,0.15055314],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.999337,0.000051899326,0.0000402925,0.00021394095,0.00027933196,0.00007743907],"domain_scores_gemma":[0.99929917,0.00014109937,0.000026063373,0.0002632197,0.0002399387,0.00003044503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00048785642,0.0008728219,0.00051738956,0.0017121296,0.0005130153,0.0019467581,0.001162142,0.000563671,0.03846046],"category_scores_gemma":[0.0031755872,0.000348173,0.0012041819,0.0014195477,0.00027149273,0.0021164245,0.0016125928,0.0012344786,0.021907398],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014120605,0.00011740763,0.001286268,0.00021586972,0.00005140648,0.0001772334,0.00009234257,0.0043382472,0.012400304,0.031124106,0.09539714,0.85465837],"study_design_scores_gemma":[0.00006730815,0.00017828857,0.0028039906,0.00014777626,0.00012457152,0.00094111904,0.00021457298,0.15635616,0.07496801,0.1330849,0.6310409,0.000072489245],"about_ca_topic_score_codex":0.002124502,"about_ca_topic_score_gemma":0.0029928316,"teacher_disagreement_score":0.03846046,"about_ca_system_score_codex":0.00058799115,"about_ca_system_score_gemma":0.00086428324,"threshold_uncertainty_score":0.12866306},"labels":[],"label_agreement":null},{"id":"W4385571400","doi":"10.18653/v1/2023.bionlp-1.5","title":"Using Bottleneck Adapters to Identify Cancer in Clinical Notes under Low-Resource Constraints","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Salud Carlos III; Canadian Institutes of Health Research; All-India Institute of Medical Sciences; Imperial College London; Conselho Nacional de Desenvolvimento Científico e Tecnológico; University of Cape Town; Public Health England; Wellcome Trust; University College Dublin; Horizon 2020 Framework Programme; Prince Charles Hospital Foundation; Foreign, Commonwealth and Development Office; National Institute for Health Research Health Protection Research Unit; University of Oxford; Norges Forskningsråd; Medical Research Council; University of Liverpool; European Federation of Pharmaceutical Industries and Associations; Health Research Board; National Institute for Health and Care Research; Ministero della Salute; Institut National de la Santé et de la Recherche Médicale; Sunnybrook Research Institute; Australian Research Council; Firland Foundation; National Institutes of Health; Kementerian Kesihatan Malaysia; European Commission; Bill and Melinda Gates Foundation","keywords":"Bottleneck; Computer science; Resource (disambiguation); Biomedical text mining; Natural language processing; Text mining; Computer network; Embedded system","score_opus":0.12039216127064416,"score_gpt":0.4453725005894047,"score_spread":0.3249803393187605,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385571400","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.54427505,0.0035945298,0.32323554,0.006030197,0.0009998268,0.0012143535,0.03682352,0.072528474,0.011298497],"genre_scores_gemma":[0.6514293,0.0009453085,0.27703378,0.0010152401,0.00027093847,0.00069162477,0.06325942,0.0014258017,0.003928591],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9966807,0.000849054,0.0004988928,0.000675356,0.0008934174,0.0004026502],"domain_scores_gemma":[0.98353875,0.007403804,0.0013612246,0.0028189295,0.0039809938,0.00089630554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0067851646,0.0012042444,0.0012393768,0.004618231,0.0012787665,0.0026135049,0.002447784,0.0014800817,0.004090435],"category_scores_gemma":[0.027893992,0.00071990816,0.0010587616,0.004113471,0.0005519615,0.005306486,0.0043245913,0.001245721,0.002293025],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.004295131,0.001468681,0.1386831,0.0013034499,0.0007603085,0.0018635938,0.0021993627,0.03413539,0.029926986,0.013178155,0.12512799,0.6470579],"study_design_scores_gemma":[0.0004356767,0.00049689476,0.030589782,0.0002655432,0.00036868113,0.00067778386,0.0032699136,0.879283,0.021857776,0.034580514,0.028038308,0.00013604008],"about_ca_topic_score_codex":0.01676666,"about_ca_topic_score_gemma":0.023366744,"teacher_disagreement_score":0.01676666,"about_ca_system_score_codex":0.0015296843,"about_ca_system_score_gemma":0.0032655986,"threshold_uncertainty_score":0.035883844},"labels":[],"label_agreement":null},{"id":"W4385573833","doi":"10.18653/v1/2022.louhi-1.25","title":"Integration of Heterogeneous Knowledge Sources for Biomedical Text Processing","year":2022,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Leverage (statistics); Interoperability; Computer science; Knowledge graph; Graph; Task (project management); Knowledge integration; Artificial intelligence; Domain knowledge; Data science; Theoretical computer science; World Wide Web; Engineering; Systems engineering","score_opus":0.024375997921282273,"score_gpt":0.3039290350910197,"score_spread":0.27955303716973745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385573833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023634652,0.001215055,0.9645353,0.0007956383,0.000115801624,0.00021375914,0.0010569996,0.0051497538,0.0032830352],"genre_scores_gemma":[0.35337472,0.0015657516,0.6324855,0.0005905871,0.00014315505,0.00032843996,0.006602126,0.0008548371,0.004054807],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982765,0.0005945387,0.00019513526,0.00041756648,0.00044043997,0.00007587564],"domain_scores_gemma":[0.9960317,0.002015094,0.00019648828,0.0009830993,0.00067153276,0.00010200456],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00267039,0.0012584277,0.0009736417,0.004569767,0.0007396952,0.001954476,0.0012314088,0.0010911837,0.00271165],"category_scores_gemma":[0.0106984805,0.0005580239,0.001448716,0.0039668097,0.0006784691,0.0059256484,0.0054366756,0.0017981046,0.0016521331],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004271117,0.00047812352,0.0050542857,0.0007787381,0.00062286726,0.0009452952,0.000729157,0.05770606,0.023595471,0.016365197,0.01038334,0.88291436],"study_design_scores_gemma":[0.000084213905,0.00019339543,0.003942497,0.0002739288,0.00054872956,0.0006758957,0.00071152416,0.76455885,0.04238799,0.14268936,0.043818433,0.0001151554],"about_ca_topic_score_codex":0.0021025804,"about_ca_topic_score_gemma":0.0048065423,"teacher_disagreement_score":0.004569767,"about_ca_system_score_codex":0.00059331954,"about_ca_system_score_gemma":0.0011169082,"threshold_uncertainty_score":0.014122546},"labels":[],"label_agreement":null},{"id":"W4385581897","doi":"10.1186/s12911-023-02250-z","title":"An ontology-based approach for harmonization and cross-cohort query of Alzheimer’s disease data resources","year":2023,"lang":"en","type":"article","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Biomedical Imaging and Bioengineering; Canadian Institutes of Health Research; National Institutes of Health; Genentech; U.S. National Library of Medicine; National Institute of Neurological Disorders and Stroke; IXICO; H. Lundbeck A/S; Servier; Eisai; Northern California Institute for Research and Education; Pfizer; Biogen; BioClinica; F. Hoffmann-La Roche; University of Southern California; Eli Lilly and Company; U.S. Department of Defense; Meso Scale Diagnostics; Alzheimer's Disease Neuroimaging Initiative; Novartis Pharmaceuticals Corporation; Bristol-Myers Squibb; National Institute on Aging; Alzheimer's Association; Foundation for the National Institutes of Health","keywords":"Ontology; Health informatics; Harmonization; Computer science; Disease; Cohort; Information retrieval; Medicine; World Wide Web; Data science; Public health; Pathology","score_opus":0.08685100647913856,"score_gpt":0.3942539625225719,"score_spread":0.3074029560434333,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385581897","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0075932946,0.00016333608,0.9744821,0.0009492949,0.00006405293,0.0010246505,0.002997772,0.010209991,0.0025154615],"genre_scores_gemma":[0.048503768,0.00019983394,0.93942,0.0004998377,0.000026101277,0.0007178359,0.008265866,0.0010947365,0.0012720142],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9801732,0.004934414,0.004841069,0.003449416,0.0059192157,0.0006827988],"domain_scores_gemma":[0.9750004,0.01053057,0.0015338942,0.0075499774,0.004592654,0.00079245603],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021969115,0.0013216435,0.001446828,0.009711145,0.0030207094,0.0060352236,0.004392105,0.0019931004,0.0030725524],"category_scores_gemma":[0.038658876,0.0017036297,0.0048840865,0.008892633,0.0023742996,0.012158853,0.011502015,0.0031399971,0.0010233505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090526824,0.0010809377,0.01999429,0.0027183273,0.0012389236,0.0037989365,0.024552854,0.04451467,0.043676976,0.28920203,0.06672842,0.50158834],"study_design_scores_gemma":[0.00031418673,0.00026723693,0.0069868076,0.0008827098,0.0007391258,0.0020974088,0.009376313,0.364527,0.040617578,0.19873033,0.37496912,0.00049225363],"about_ca_topic_score_codex":0.030631376,"about_ca_topic_score_gemma":0.03144166,"teacher_disagreement_score":0.030631376,"about_ca_system_score_codex":0.0041681,"about_ca_system_score_gemma":0.0076105245,"threshold_uncertainty_score":0.11618513},"labels":[],"label_agreement":null},{"id":"W4385709848","doi":"10.1093/jamiaopen/ooad062","title":"Automated identification of unstandardized medication data: a scalable and flexible data standardization pipeline using RxNorm on GEMINI multicenter hospital data","year":2023,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto; St. Michael's Hospital","funders":"Alliance de recherche numérique du Canada","keywords":"Standardization; Computer science; Identifier; Pharmacy; Coding (social sciences); Data mining; Information retrieval; Medicine; Statistics; Mathematics; Family medicine","score_opus":0.09741173152302603,"score_gpt":0.40517662857133674,"score_spread":0.3077648970483107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385709848","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13347626,0.001609107,0.55196005,0.0031479888,0.00017790252,0.0024162375,0.041592613,0.25943732,0.006182528],"genre_scores_gemma":[0.21470842,0.0005385302,0.71253103,0.0007375658,0.000099327444,0.0008361132,0.06495054,0.003678691,0.0019197382],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.98601365,0.0034909472,0.0022137917,0.0034546969,0.0043758885,0.00045105073],"domain_scores_gemma":[0.9682873,0.010216466,0.004793179,0.008882719,0.007109439,0.0007107822],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022917943,0.0016656471,0.001127704,0.006035862,0.0010977591,0.0034962941,0.0022568454,0.0006959593,0.0018972016],"category_scores_gemma":[0.042838734,0.0008868027,0.0015324316,0.0047239214,0.0010596693,0.004105562,0.0042527956,0.0014682435,0.0018112634],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021620223,0.0006583183,0.14018403,0.0014602554,0.0007950412,0.0011027312,0.003911375,0.020629345,0.049670346,0.00834034,0.087228775,0.6838574],"study_design_scores_gemma":[0.0006620859,0.00087631785,0.12037696,0.00066621305,0.00041343115,0.0014726399,0.0024506259,0.4801412,0.1960168,0.017175093,0.17915413,0.00059453346],"about_ca_topic_score_codex":0.019262845,"about_ca_topic_score_gemma":0.015064242,"teacher_disagreement_score":0.022917943,"about_ca_system_score_codex":0.002627307,"about_ca_system_score_gemma":0.006855585,"threshold_uncertainty_score":0.121203125},"labels":[],"label_agreement":null},{"id":"W4385827338","doi":"10.1007/s00335-023-10014-3","title":"Mouse phenome database: curated data repository with interactive multi-population and multi-trait analyses","year":2023,"lang":"en","type":"review","venue":"Mammalian Genome","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"National Cancer Institute; National Institute on Drug Abuse; National Institute on Aging; National Institutes of Health","keywords":"Phenome; Biology; Population; Database; Suite; Annotation; Computer science; Bioinformatics; Genetics; Phenotype; Gene","score_opus":0.22857827769888836,"score_gpt":0.42573051664394174,"score_spread":0.19715223894505338,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385827338","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037176537,0.4493324,0.107660286,0.0076142037,0.0026806688,0.0015293085,0.3540648,0.033850964,0.039549798],"genre_scores_gemma":[0.009365816,0.35408226,0.10405001,0.004758249,0.0010533154,0.002287084,0.5084763,0.0033379903,0.012588989],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986621,0.0002488906,0.00023006875,0.00018239276,0.00057917106,0.00009724126],"domain_scores_gemma":[0.99652773,0.0010413056,0.00040236657,0.00046623108,0.0011573363,0.00040507535],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003562973,0.0013649227,0.002963946,0.009624449,0.0004610665,0.0025999362,0.004705806,0.001289448,0.0169513],"category_scores_gemma":[0.005053931,0.0005392087,0.0016517328,0.008965579,0.00042662004,0.002064229,0.0024988984,0.0022279741,0.012267832],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024839898,0.000092970564,0.00067383633,0.015335066,0.00064758956,0.0003295643,0.00013881933,0.00081573473,0.0048576957,0.010446031,0.46544543,0.5009689],"study_design_scores_gemma":[0.000043850523,0.00001727815,0.0012898635,0.0016204756,0.00017878314,0.00036497542,0.000024958199,0.00020849607,0.001075096,0.0030527532,0.992084,0.00003956324],"about_ca_topic_score_codex":0.002154769,"about_ca_topic_score_gemma":0.00324646,"teacher_disagreement_score":0.0169513,"about_ca_system_score_codex":0.00088565884,"about_ca_system_score_gemma":0.003957424,"threshold_uncertainty_score":0.05670774},"labels":[],"label_agreement":null},{"id":"W4385848332","doi":"10.1016/j.patter.2023.100887","title":"Enhancing phenotype recognition in clinical notes using large language models: PhenoBCBERT and PhenoGPT","year":2023,"lang":"en","type":"article","venue":"Patterns","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"U.S. National Library of Medicine; National Human Genome Research Institute; Intellectual and Developmental Disabilities Research Center; CHEO Research Institute; University of Pennsylvania; National Institutes of Health; Eunice Kennedy Shriver National Institute of Child Health and Human Development; Children's Hospital of Philadelphia","keywords":"Phenotype; Leverage (statistics); Computer science; Vocabulary; Heuristic; Ontology; Natural language processing; Scope (computer science); Artificial intelligence; Computational biology; Machine learning; Biology; Gene; Genetics; Linguistics","score_opus":0.07803578452570843,"score_gpt":0.3586311829778933,"score_spread":0.2805953984521849,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385848332","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09563865,0.0010362307,0.8466375,0.004676316,0.0003088176,0.0008939957,0.015333188,0.030965347,0.0045100003],"genre_scores_gemma":[0.38047767,0.0006613707,0.59620386,0.001473864,0.00012132852,0.00061573467,0.01692276,0.0007938981,0.0027295141],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99844885,0.00047465856,0.0002003932,0.00038538087,0.00043287946,0.0000578162],"domain_scores_gemma":[0.99184704,0.005880296,0.0005293118,0.0005851997,0.00091603637,0.00024205339],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027973559,0.00092833926,0.0006619075,0.002259144,0.00034895074,0.0017623648,0.0011316604,0.0009399331,0.0019087311],"category_scores_gemma":[0.011491334,0.0003823699,0.001346093,0.001136544,0.00045282234,0.0023017589,0.0014136523,0.00141011,0.000941752],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011602226,0.00085360295,0.051441643,0.0012527959,0.0005712566,0.0023671687,0.0010100569,0.33067015,0.021122046,0.021958416,0.048268635,0.51932406],"study_design_scores_gemma":[0.00006157662,0.000060961807,0.0024832992,0.000065252825,0.000070186165,0.00034320416,0.00008687903,0.97209394,0.0052486705,0.010288038,0.009153128,0.000044836495],"about_ca_topic_score_codex":0.014018858,"about_ca_topic_score_gemma":0.026130162,"teacher_disagreement_score":0.014018858,"about_ca_system_score_codex":0.0013145944,"about_ca_system_score_gemma":0.0023363926,"threshold_uncertainty_score":0.02787453},"labels":[],"label_agreement":null},{"id":"W4385893888","doi":"10.18653/v1/2023.bionlp-1.25","title":"Extracting Drug-Drug and Protein-Protein Interactions from Text using a Continuous Update of Tree-Transformers","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Sentence; Computer science; Transformer; Artificial intelligence; Natural language processing; Tree (set theory); Mathematics","score_opus":0.017713114628669483,"score_gpt":0.28489940858510693,"score_spread":0.26718629395643745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4385893888","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06885529,0.0017593817,0.9036408,0.0016531487,0.00033354975,0.00047585214,0.007475786,0.0112787,0.0045274533],"genre_scores_gemma":[0.506801,0.0017382855,0.46693102,0.00050362386,0.00026793842,0.00046866285,0.016982157,0.00059165875,0.0057156435],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9995185,0.00008860901,0.000058181082,0.00018402548,0.000119096105,0.00003156139],"domain_scores_gemma":[0.99843735,0.00092288444,0.00013405159,0.00013487357,0.0003140663,0.000056762146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00070199446,0.0011869769,0.00056937797,0.0020540974,0.0003093719,0.0010180654,0.0011585591,0.00089024927,0.002579552],"category_scores_gemma":[0.0036847165,0.0004617797,0.0013909434,0.0017192513,0.00041227273,0.0032188625,0.00090206065,0.0013323043,0.0022894235],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007477162,0.00036002425,0.00742916,0.0008445531,0.00020357399,0.001234979,0.00066672557,0.10809338,0.040690113,0.014386976,0.02647014,0.79887265],"study_design_scores_gemma":[0.000041676976,0.00013829986,0.0017151306,0.0000439502,0.00012913739,0.00039240922,0.00011312955,0.95688784,0.011543937,0.018121447,0.010838818,0.000034225184],"about_ca_topic_score_codex":0.004096543,"about_ca_topic_score_gemma":0.008401112,"teacher_disagreement_score":0.004096543,"about_ca_system_score_codex":0.0007049439,"about_ca_system_score_gemma":0.0014097033,"threshold_uncertainty_score":0.008629441},"labels":[],"label_agreement":null},{"id":"W4386090599","doi":"10.1093/zoolinnean/zlad056","title":"Governance of biological nomenclature: mechanisms to address the needs of end-users are available and not onerous to implement","year":2023,"lang":"en","type":"article","venue":"Zoological Journal of the Linnean Society","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Winnipeg","funders":"Linnean Society of London","keywords":"Biology; Nomenclature; Corporate governance; Evolutionary biology; Zoology; Taxonomy (biology); Management","score_opus":0.039845527960618056,"score_gpt":0.2790462568326274,"score_spread":0.23920072887200935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386090599","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019194536,0.0073585683,0.36935705,0.31756943,0.0051459293,0.0015984215,0.00033518317,0.0026919937,0.27674893],"genre_scores_gemma":[0.5675647,0.0067224177,0.21102847,0.092774734,0.0067308303,0.005849105,0.0012204493,0.0026786877,0.10543062],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.88870066,0.063513465,0.009128129,0.009738421,0.02310787,0.0058114897],"domain_scores_gemma":[0.8338803,0.044906154,0.019395888,0.044841103,0.040556204,0.016420474],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.14575864,0.0008236906,0.0016036518,0.0039774785,0.010263336,0.027220171,0.0064536403,0.010588402,0.009010038],"category_scores_gemma":[0.120212376,0.0012333656,0.00091132685,0.004814653,0.033528402,0.0406328,0.028514853,0.013766099,0.008705749],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040061306,0.0001085838,0.006126782,0.00041721974,0.000049748003,0.00027954418,0.035052903,0.00075355347,0.002132988,0.7168682,0.10134414,0.13682619],"study_design_scores_gemma":[0.000027987982,0.00005341858,0.0023009493,0.00077044143,0.000017784781,0.00019268843,0.008006519,0.0008688574,0.00056445936,0.28713605,0.6999557,0.00010515999],"about_ca_topic_score_codex":0.0036496264,"about_ca_topic_score_gemma":0.0037415063,"teacher_disagreement_score":0.14575864,"about_ca_system_score_codex":0.009270886,"about_ca_system_score_gemma":0.035343118,"threshold_uncertainty_score":0.7708546},"labels":[],"label_agreement":null},{"id":"W4386215146","doi":"10.1109/icdh60066.2023.00026","title":"SPaDe: A Synonym-based Pain-level Detection Tool for Osteoarthritis","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"Pfizer","keywords":"Osteoarthritis; Synonym (taxonomy); Computer science; Artificial intelligence; Natural language processing; Medicine; Alternative medicine; Biology; Pathology","score_opus":0.029703500311154065,"score_gpt":0.2745914642083862,"score_spread":0.24488796389723214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386215146","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21999817,0.0052511324,0.59580415,0.0045539276,0.00068508805,0.003441375,0.12606916,0.026847005,0.01735002],"genre_scores_gemma":[0.33322027,0.0014802582,0.59938496,0.0010291485,0.0001849506,0.0012849073,0.058520302,0.0004021962,0.0044929953],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9983639,0.00025812048,0.000364723,0.00033327797,0.0006026344,0.0000771717],"domain_scores_gemma":[0.99598455,0.0018274818,0.00096902926,0.00029150193,0.0007076904,0.00021976331],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001336961,0.0007253316,0.00052157417,0.006019683,0.0005403401,0.0012167844,0.0006373619,0.000828561,0.004231981],"category_scores_gemma":[0.008239892,0.00022116805,0.0008699487,0.0022492914,0.0002832456,0.002240754,0.0019628082,0.000759688,0.0017270895],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010521447,0.0007519128,0.08817039,0.0048898975,0.0005384864,0.0029158813,0.0028577056,0.002972959,0.06458193,0.016450146,0.10112987,0.7136887],"study_design_scores_gemma":[0.00066086045,0.0015853898,0.2925018,0.0015123007,0.000761346,0.014813012,0.0063629015,0.1650799,0.071746305,0.06516363,0.37930474,0.0005077863],"about_ca_topic_score_codex":0.0014537042,"about_ca_topic_score_gemma":0.00332353,"teacher_disagreement_score":0.006019683,"about_ca_system_score_codex":0.00054708414,"about_ca_system_score_gemma":0.0012879365,"threshold_uncertainty_score":0.014157414},"labels":[],"label_agreement":null},{"id":"W4386307955","doi":"10.3389/fninf.2023.1215261","title":"NeuroBridge: a prototype platform for discovery of the long-tail neuroimaging data","year":2023,"lang":"en","type":"article","venue":"Frontiers in Neuroinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Division of Information and Intelligent Systems; National Institute of Biomedical Imaging and Bioengineering; National Institute of Mental Health; Fondation Brain Canada; McGill University; Michael J. Fox Foundation for Parkinson's Research; National Institute on Drug Abuse; Health Canada; Canada First Research Excellence Fund; National Institutes of Health; National Science Foundation","keywords":"Computer science; Ontology; Information retrieval; Neuroinformatics; Metadata; Neuroimaging; Data science; World Wide Web; Artificial intelligence; Natural language processing; Psychology","score_opus":0.039271864788121696,"score_gpt":0.2874974096432426,"score_spread":0.2482255448551209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386307955","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022299808,0.003013473,0.22749506,0.0042591398,0.00058514136,0.0019657274,0.2663658,0.4463626,0.027653344],"genre_scores_gemma":[0.07977565,0.002242022,0.3989851,0.0026426066,0.00027583272,0.0034794346,0.46989223,0.0205619,0.022145241],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99822205,0.00028989703,0.00022565259,0.00051336864,0.00063221285,0.00011687002],"domain_scores_gemma":[0.994179,0.002822149,0.00047216308,0.0012664627,0.0007428668,0.00051741843],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.0041555315,0.0016138148,0.00097224384,0.0071041803,0.00089658215,0.004066303,0.0035278962,0.0014866706,0.026268948],"category_scores_gemma":[0.014127773,0.00091951515,0.0013824899,0.0048903315,0.0008847223,0.0071912613,0.0060463697,0.0013439634,0.018265273],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023822263,0.00036244624,0.007132414,0.0040005916,0.0005095882,0.0016468356,0.0017143541,0.0036952312,0.029563697,0.022307072,0.73178923,0.19489636],"study_design_scores_gemma":[0.0007569384,0.00039309615,0.0067784223,0.00054552715,0.00016789928,0.001089745,0.0010016747,0.06389745,0.024812708,0.042376257,0.8578984,0.000281909],"about_ca_topic_score_codex":0.0071590524,"about_ca_topic_score_gemma":0.016962754,"teacher_disagreement_score":0.9964721,"about_ca_system_score_codex":0.0015182189,"about_ca_system_score_gemma":0.0030422532,"threshold_uncertainty_score":0.087878406},"labels":[],"label_agreement":null},{"id":"W4386466998","doi":"10.23889/ijpds.v6i1.1757","title":"A scoping review of preprocessing methods for unstructured text data to assess data quality.","year":2022,"lang":"en","type":"review","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; George & Fay Yee Centre for Healthcare Innovation; Manitoba Health","funders":"Canada Research Chairs","keywords":"Computer science; Preprocessor; Punctuation; Stop words; Data quality; Data pre-processing; Information retrieval; Lexical analysis; Natural language processing; Artificial intelligence; Data mining","score_opus":0.6516293020376296,"score_gpt":0.5791308034657566,"score_spread":0.07249849857187296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386466998","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015921902,0.93671983,0.02041633,0.007063176,0.002319143,0.024886448,0.0031439648,0.00025128233,0.0036076137],"genre_scores_gemma":[0.015114828,0.8142836,0.075438626,0.0035194545,0.00070320145,0.0859601,0.0037286696,0.0002208503,0.0010306281],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.67910963,0.13433307,0.1303153,0.00815331,0.046435263,0.0016534403],"domain_scores_gemma":[0.28463912,0.51348597,0.074043564,0.017789628,0.10798214,0.002059633],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.32512677,0.0038456605,0.011046392,0.058586884,0.005064684,0.0130500095,0.0075221206,0.0065544895,0.00848936],"category_scores_gemma":[0.58527476,0.0038252703,0.014160485,0.04518956,0.0065765344,0.01167418,0.009601528,0.004953821,0.0024511968],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024843094,0.000053069787,0.001420572,0.7597259,0.0029033376,0.00017516412,0.0029450434,0.00031294487,0.00046817245,0.0027472733,0.010764466,0.21823561],"study_design_scores_gemma":[0.00007535446,0.0000927323,0.0011949154,0.95408267,0.0037745116,0.00015836896,0.0008207847,0.00016448445,0.000406351,0.0016270775,0.037539974,0.00006263844],"about_ca_topic_score_codex":0.01017957,"about_ca_topic_score_gemma":0.018432181,"teacher_disagreement_score":0.67487323,"about_ca_system_score_codex":0.018446036,"about_ca_system_score_gemma":0.07373671,"threshold_uncertainty_score":0.83223885},"labels":[],"label_agreement":null},{"id":"W4386645083","doi":"10.18280/isi.280424","title":"A Multi-Agent Systems Approach for Optimized Biomedical Literature Search","year":2023,"lang":"fr","type":"article","venue":"Ingénierie des systèmes d information","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Management science; Data science; Engineering","score_opus":0.03422154045463785,"score_gpt":0.2881653716047316,"score_spread":0.2539438311500938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386645083","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048149913,0.0018652076,0.9855203,0.0008195709,0.00009548992,0.0002590682,0.00008740655,0.00043105663,0.0061068223],"genre_scores_gemma":[0.24978644,0.0020711038,0.7407069,0.00030772292,0.00014648453,0.00078968116,0.00025143797,0.0000934058,0.0058467183],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985362,0.00069129735,0.00011376863,0.00025104053,0.00032821545,0.00007944232],"domain_scores_gemma":[0.99856985,0.0007826995,0.00017125459,0.000077545024,0.00030244165,0.00009617264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024526406,0.000961736,0.0012368637,0.0020295938,0.0011575426,0.0029142185,0.0017430269,0.0015229187,0.0034108595],"category_scores_gemma":[0.0053694323,0.0005792072,0.0010057449,0.0019146557,0.0007610607,0.0016069233,0.0023922317,0.0011623766,0.0007321753],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000102626975,0.000086802436,0.00086903863,0.0006015302,0.00023400523,0.00043398782,0.0004364139,0.77738607,0.0033800793,0.07202554,0.005166865,0.13927704],"study_design_scores_gemma":[0.00002563255,0.00005222696,0.00012634833,0.000050426024,0.000043523047,0.000077327604,0.000090811256,0.96536916,0.00061860256,0.025637822,0.007887344,0.000020792722],"about_ca_topic_score_codex":0.0058871885,"about_ca_topic_score_gemma":0.008254373,"teacher_disagreement_score":0.0058871885,"about_ca_system_score_codex":0.0017382577,"about_ca_system_score_gemma":0.004033859,"threshold_uncertainty_score":0.012970984},"labels":[],"label_agreement":null},{"id":"W4386742622","doi":"10.1007/s41666-023-00146-1","title":"Sequence Labeling for Disambiguating Medical Abbreviations","year":2023,"lang":"en","type":"article","venue":"Journal of Healthcare Informatics Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Toronto","funders":"","keywords":"Computer science; Sequence labeling; Task (project management); Natural language processing; Sequence (biology); Artificial intelligence; Information retrieval; Transformer","score_opus":0.28723082685554796,"score_gpt":0.5417832293555354,"score_spread":0.2545524024999874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386742622","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.046543796,0.0022222772,0.87986624,0.0017458746,0.0007874435,0.00081649807,0.033571284,0.023608794,0.01083783],"genre_scores_gemma":[0.10051767,0.0006477458,0.85496306,0.00038060566,0.00014555008,0.00034275965,0.038862377,0.00070544804,0.0034348592],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99871993,0.00029566293,0.00025472234,0.00038582738,0.00026200482,0.00008183419],"domain_scores_gemma":[0.9954867,0.00233132,0.0003988285,0.0005616022,0.0010041716,0.0002173426],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012805478,0.000957433,0.00081633776,0.007125867,0.0015465807,0.0015928638,0.0011468513,0.0013579568,0.010452541],"category_scores_gemma":[0.006331123,0.00043616828,0.0010635538,0.0050027897,0.00049641693,0.0027071238,0.0015402774,0.0012749721,0.0061174184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00075423904,0.0003041484,0.006988045,0.0013388734,0.00010068055,0.0009923659,0.0011267759,0.006222977,0.043729592,0.03791338,0.078647755,0.8218813],"study_design_scores_gemma":[0.00026050626,0.0003921874,0.008117798,0.0012664765,0.00051565736,0.00281951,0.0024176752,0.3575784,0.09209921,0.14324212,0.39108363,0.00020690529],"about_ca_topic_score_codex":0.004354554,"about_ca_topic_score_gemma":0.006651388,"teacher_disagreement_score":0.010452541,"about_ca_system_score_codex":0.0010341987,"about_ca_system_score_gemma":0.0031857225,"threshold_uncertainty_score":0.034967244},"labels":[],"label_agreement":null},{"id":"W4386905201","doi":"10.1016/j.ijrobp.2023.02.058","title":"Operational Ontology for Oncology: A Framework for Improved Communication and Understanding in Cancer Care","year":2023,"lang":"en","type":"letter","venue":"International Journal of Radiation Oncology*Biology*Physics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Dalhousie University","funders":"","keywords":"Medicine; Standardization; Multidisciplinary approach; Ontology; CLARITY; Medical education; Oncology; Computer science","score_opus":0.06564763591964595,"score_gpt":0.39955920549548374,"score_spread":0.3339115695758378,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4386905201","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010307337,0.0021557705,0.12110305,0.8626234,0.0054507637,0.0001022454,0.0005771747,0.0006588845,0.006298033],"genre_scores_gemma":[0.07625804,0.0087939305,0.5471515,0.32763746,0.019500123,0.0009031684,0.0026885737,0.00076778245,0.016299449],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9896979,0.004672376,0.001533126,0.00056582486,0.0030248926,0.00050588115],"domain_scores_gemma":[0.9478966,0.03203458,0.0022734017,0.005101343,0.010097683,0.0025963325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022149991,0.00052962435,0.000996875,0.0021876602,0.0028032984,0.006800852,0.0032647683,0.008406182,0.005008616],"category_scores_gemma":[0.04464202,0.00050326216,0.0017737234,0.0023035302,0.006687347,0.016222471,0.006567916,0.015863307,0.0020616956],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000077993616,0.000076446435,0.0013550002,0.00043128294,0.00008056992,0.00059341773,0.0016579239,0.0014786191,0.0009708995,0.45262927,0.42036468,0.12028389],"study_design_scores_gemma":[0.000040395127,0.000022060754,0.00046628586,0.00035974302,0.000041034622,0.00064542925,0.00076956616,0.0054205186,0.00051493815,0.3546545,0.63701314,0.00005250706],"about_ca_topic_score_codex":0.013383691,"about_ca_topic_score_gemma":0.018215496,"teacher_disagreement_score":0.022149991,"about_ca_system_score_codex":0.009309845,"about_ca_system_score_gemma":0.016179515,"threshold_uncertainty_score":0.11714178},"labels":[],"label_agreement":null},{"id":"W4387357939","doi":"10.2196/44892","title":"A Multilabel Text Classifier of Cancer Literature at the Publication Level: Methods Study of Medical Text Classification","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Classifier (UML); Artificial intelligence; Machine learning; Terminology; Information retrieval; Natural language processing","score_opus":0.08301918992268613,"score_gpt":0.43616050301028353,"score_spread":0.3531413130875974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387357939","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4362189,0.005681768,0.54231256,0.0026095847,0.00051844754,0.00074137305,0.0042328667,0.0027824633,0.0049020234],"genre_scores_gemma":[0.79866844,0.0006396916,0.18967956,0.0003017879,0.00043921688,0.0004372509,0.0054992163,0.00011498279,0.0042198026],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964336,0.0013029199,0.00034997973,0.0009119182,0.0007559195,0.00024569238],"domain_scores_gemma":[0.9783327,0.01572927,0.0014728203,0.0014195334,0.0025505256,0.0004950934],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0071039554,0.00082850334,0.0009354317,0.005038618,0.0007671016,0.001818099,0.0015273091,0.001591585,0.0018857042],"category_scores_gemma":[0.019640712,0.00021134944,0.0011012496,0.0030228214,0.0008448742,0.0028722,0.001284595,0.0018284668,0.0011799285],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016484248,0.0009625068,0.06848899,0.0009870578,0.00048521548,0.00047062928,0.00067854114,0.055609312,0.010635422,0.007705152,0.012456838,0.8398718],"study_design_scores_gemma":[0.00006879075,0.000313492,0.010686494,0.00010605899,0.0001540069,0.0003175345,0.0002964128,0.9670515,0.007833164,0.009527321,0.003600942,0.000044291348],"about_ca_topic_score_codex":0.0034627027,"about_ca_topic_score_gemma":0.0032040013,"teacher_disagreement_score":0.992896,"about_ca_system_score_codex":0.0016282572,"about_ca_system_score_gemma":0.0018988313,"threshold_uncertainty_score":0.0375697},"labels":[],"label_agreement":null},{"id":"W4387741944","doi":"10.1111/mila.12475","title":"Who's in and who's out of the cognitive kinding game? Comments on Muhammad Ali Khalidi's <i>Cognitive ontology: Taxonomic practices in the mind‐brain sciences</i>","year":2023,"lang":"en","type":"article","venue":"Mind & Language","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Cognition; Psychology; Cognitive neuroscience; Cognitive science; Ontology; Embodied cognition; Epistemology; Animal cognition; Sociology; Cognitive psychology; Philosophy; Neuroscience","score_opus":0.06202439634147285,"score_gpt":0.36616468308459177,"score_spread":0.30414028674311894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387741944","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00052147795,0.00097287045,0.00012281442,0.994212,0.0024454726,0.0000036797048,0.00002339314,0.0000142836925,0.0016838871],"genre_scores_gemma":[0.018484166,0.001385612,0.0003503379,0.9674982,0.0060160663,0.000042537384,0.00001862483,0.00004671336,0.0061577437],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9948166,0.0018921457,0.00034195185,0.0010051895,0.001217369,0.0007268681],"domain_scores_gemma":[0.9700551,0.021015035,0.0015665168,0.00061320286,0.0048479605,0.0019020886],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.012569097,0.00089134503,0.0012500029,0.0008443644,0.011159236,0.008067999,0.0043424396,0.027417798,0.0057209954],"category_scores_gemma":[0.029736986,0.00054240826,0.0008708971,0.0012604306,0.015978882,0.0137106255,0.004013435,0.03755773,0.0022918994],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003687041,0.000017514452,0.00037413248,0.00005142318,0.000007348535,0.00048614317,0.008915992,0.00006434485,0.000111234556,0.041993,0.94338405,0.004558019],"study_design_scores_gemma":[0.000024177654,0.000026274274,0.0012759073,0.00034301888,0.00001662405,0.00060135964,0.03677989,0.00032050713,0.00034099063,0.027427077,0.9326988,0.000145453],"about_ca_topic_score_codex":0.048915002,"about_ca_topic_score_gemma":0.046667997,"teacher_disagreement_score":0.98884076,"about_ca_system_score_codex":0.0071896715,"about_ca_system_score_gemma":0.006387119,"threshold_uncertainty_score":0.097260535},"labels":[],"label_agreement":null},{"id":"W4387781486","doi":"10.1186/s12911-023-02333-x","title":"Correction: Development and validation of a case definition for problematic menopause in primary care electronic medical records","year":2023,"lang":"en","type":"erratum","venue":"BMC Medical Informatics and Decision Making","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; University of Alberta","funders":"","keywords":"Health informatics; Primary care; Medical record; Medicine; Computer science; Data science; Family medicine; Nursing; Public health","score_opus":0.03395377599764433,"score_gpt":0.3152371215469454,"score_spread":0.28128334554930107,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387781486","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000951321,0.00095097494,0.0052337656,0.2152734,0.7649579,0.0002229125,0.0060043884,0.0008698341,0.0055354573],"genre_scores_gemma":[0.053866755,0.010288341,0.06404601,0.32541505,0.27458775,0.0022025034,0.011820933,0.0047279107,0.2530448],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98811054,0.0026303902,0.0031088595,0.0012069156,0.0043241037,0.0006190996],"domain_scores_gemma":[0.8626712,0.054020654,0.006363091,0.0065155285,0.067656994,0.002772519],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012129451,0.0014912296,0.0015324273,0.004654125,0.0032460822,0.0037603925,0.004374014,0.008398761,0.035154656],"category_scores_gemma":[0.2039355,0.0011092272,0.0013143433,0.0028984135,0.0037144262,0.0024790107,0.003451722,0.0104217995,0.017073302],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003148163,0.000006655797,0.0002188923,0.00024422718,0.00001541389,0.0008528741,0.00022550925,0.000049178558,0.00006883037,0.0009874058,0.9908249,0.00647462],"study_design_scores_gemma":[0.000081249236,0.000020978066,0.0016678281,0.0013962436,0.00008715456,0.0023804612,0.00047933895,0.00047396054,0.00057704514,0.0021877072,0.9905772,0.000070860355],"about_ca_topic_score_codex":0.03384982,"about_ca_topic_score_gemma":0.03550028,"teacher_disagreement_score":0.035154656,"about_ca_system_score_codex":0.0050805984,"about_ca_system_score_gemma":0.012117265,"threshold_uncertainty_score":0.11760408},"labels":[],"label_agreement":null},{"id":"W4387952669","doi":"10.20944/preprints202310.1692.v1","title":"Cardiovascular Disease Preliminary Diagnosis Application Using SQL Queries: Filling Diagnostic Gaps in&#x0D;Resource-Constrained Environments","year":2023,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec en Outaouais","funders":"Universität zu Lübeck","keywords":"Disease; Psychological intervention; Ontology; Computer science; Resource (disambiguation); Medicine; SQL; Intensive care medicine; Database; Pathology; Nursing","score_opus":0.07674474170057873,"score_gpt":0.3131366625500906,"score_spread":0.23639192084951188,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4387952669","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13653189,0.002317357,0.5684716,0.006480636,0.0005627782,0.0012993284,0.038850997,0.22731857,0.018166853],"genre_scores_gemma":[0.44374362,0.0015713964,0.49056372,0.0023831017,0.00020696966,0.0006065693,0.05009169,0.0044956827,0.0063372804],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99769044,0.00048883096,0.0004452699,0.00059630966,0.00061700866,0.0001621318],"domain_scores_gemma":[0.9928104,0.0044509633,0.00041261088,0.000998669,0.0008720665,0.00045532957],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036556185,0.001148301,0.0009787943,0.0023762623,0.00091574027,0.0040637157,0.002296881,0.0018418622,0.0088346945],"category_scores_gemma":[0.012700589,0.00063683395,0.0009693136,0.001940312,0.0006612165,0.004394345,0.0042325663,0.0010567387,0.0032095965],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0049947873,0.0015873665,0.08068674,0.0033572451,0.00048228583,0.008191938,0.0066342736,0.03602439,0.039577525,0.048970684,0.19098064,0.57851213],"study_design_scores_gemma":[0.00068383367,0.00041926763,0.016564922,0.0005780868,0.00024576118,0.0029659036,0.0042741504,0.6482419,0.04664693,0.055243127,0.22390762,0.00022841775],"about_ca_topic_score_codex":0.006795575,"about_ca_topic_score_gemma":0.006147827,"teacher_disagreement_score":0.0088346945,"about_ca_system_score_codex":0.00089016184,"about_ca_system_score_gemma":0.0017432483,"threshold_uncertainty_score":0.029554963},"labels":[],"label_agreement":null},{"id":"W4388041995","doi":"10.1016/j.cancergen.2023.08.025","title":"17. Djerba: A modular system to generate clinical genome interpretation reports for cancer","year":2023,"lang":"en","type":"article","venue":"Cancer Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"","keywords":"Modular design; Interpretation (philosophy); Computational biology; Computer science; Cancer; Biology; Genetics; Programming language","score_opus":0.04554806236866749,"score_gpt":0.38012802717262156,"score_spread":0.3345799648039541,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388041995","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005442726,0.0006780032,0.2736872,0.0013154087,0.00025221836,0.0004979794,0.09263829,0.61868167,0.0068065394],"genre_scores_gemma":[0.096022025,0.0015669384,0.47821733,0.0024209055,0.0002557942,0.0014219228,0.32437292,0.07267612,0.023046033],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99917126,0.00017419423,0.00012160136,0.00021781775,0.00024520126,0.00006986243],"domain_scores_gemma":[0.9966912,0.0018148471,0.00031231498,0.0004850129,0.0005126053,0.00018394894],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033631155,0.0014945726,0.00096790586,0.0049211956,0.0006692717,0.0029997376,0.0018529042,0.0013875385,0.04370552],"category_scores_gemma":[0.009328749,0.0013374838,0.0013564348,0.0022068196,0.0005272942,0.002070133,0.00279145,0.0011542223,0.021054422],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025799866,0.0003083458,0.017944818,0.0039685303,0.0008621935,0.0015643165,0.0011237168,0.0065188967,0.02551173,0.012674566,0.6431065,0.2838363],"study_design_scores_gemma":[0.0013696989,0.00020821246,0.012953087,0.00057716516,0.00053699705,0.0019277958,0.00031908773,0.07578813,0.07687024,0.023675097,0.80543154,0.0003429262],"about_ca_topic_score_codex":0.0051358477,"about_ca_topic_score_gemma":0.0071786176,"teacher_disagreement_score":0.04370552,"about_ca_system_score_codex":0.001180734,"about_ca_system_score_gemma":0.0020591356,"threshold_uncertainty_score":0.14620954},"labels":[],"label_agreement":null},{"id":"W4388089782","doi":"10.1109/isc257844.2023.10293563","title":"A New Semantic Similarity Scheme for More Accurate Identification in Medical Data","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Jaccard index; Computer science; Similarity (geometry); Semantic similarity; Context (archaeology); Identification (biology); Information retrieval; Cosine similarity; Data mining; Set (abstract data type); Data set; Dice; Key (lock); Similarity measure; Scheme (mathematics); Artificial intelligence; Pattern recognition (psychology); Mathematics; Statistics","score_opus":0.08753904631601674,"score_gpt":0.3982877170484738,"score_spread":0.3107486707324571,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388089782","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03351116,0.00053317857,0.96234596,0.00030577218,0.00013476724,0.00028967214,0.0004746681,0.00054645893,0.0018583325],"genre_scores_gemma":[0.25983188,0.0002757482,0.7372108,0.00013377605,0.00010607272,0.00027515794,0.0010854152,0.00007036157,0.0010106869],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99326247,0.0015263013,0.0013775248,0.0008733538,0.0027684965,0.00019191405],"domain_scores_gemma":[0.9929651,0.0021508273,0.0010920453,0.0014682896,0.0020361973,0.00028743903],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005108262,0.0004255859,0.0008598476,0.006142169,0.00080688205,0.0017850447,0.0012716613,0.0008965484,0.00157562],"category_scores_gemma":[0.018640094,0.00018408694,0.0007745333,0.0052659055,0.0010826528,0.0056602145,0.0027875302,0.00096218684,0.0008907048],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000687808,0.00029753172,0.014676987,0.0009471559,0.00024740503,0.0002931931,0.0017497102,0.021728009,0.05129191,0.12147176,0.0060158772,0.7805926],"study_design_scores_gemma":[0.00015235977,0.001918128,0.02283325,0.0005033801,0.0002768502,0.0042729606,0.0020660951,0.6016837,0.07708016,0.21016142,0.0786623,0.00038940844],"about_ca_topic_score_codex":0.0007011623,"about_ca_topic_score_gemma":0.0007277303,"teacher_disagreement_score":0.006142169,"about_ca_system_score_codex":0.0010855265,"about_ca_system_score_gemma":0.0013753513,"threshold_uncertainty_score":0.027015448},"labels":[],"label_agreement":null},{"id":"W4388142096","doi":"10.1371/journal.pdig.0000388","title":"Hallucination or Confabulation? Neuroanatomy as metaphor in Large Language Models","year":2023,"lang":"en","type":"article","venue":"PLOS Digital Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":73,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; University of Ottawa","funders":"","keywords":"Confabulation (neural networks); Metaphor; Psychology; Cognitive psychology; Cognitive science; Linguistics; Cognition; Neuroscience; Philosophy","score_opus":0.0338593488114151,"score_gpt":0.334673976730809,"score_spread":0.3008146279193939,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388142096","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13596494,0.0013782696,0.8364721,0.011765787,0.00019979302,0.00009637672,0.0013902922,0.0015360055,0.011196456],"genre_scores_gemma":[0.9042272,0.00037891752,0.092455015,0.00049179257,0.000078752484,0.00010979299,0.00049971987,0.00015063632,0.0016081027],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9987956,0.0008112317,0.00005682125,0.00017672047,0.00011626231,0.000043377935],"domain_scores_gemma":[0.989239,0.00912954,0.00040879817,0.0007436334,0.00026825292,0.00021075743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024426873,0.0003963758,0.00054896536,0.0010211004,0.00060652994,0.002536572,0.0010576113,0.0009390948,0.004330215],"category_scores_gemma":[0.021439502,0.00038200556,0.000874219,0.0010964532,0.0025988966,0.008481481,0.0019895658,0.0019276022,0.00046558882],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041691432,0.0001288855,0.009266116,0.00057299953,0.00025980326,0.0013113238,0.004775741,0.05217111,0.004746794,0.7644142,0.011262862,0.15067329],"study_design_scores_gemma":[0.00002309658,0.000025398802,0.00069388864,0.00004175933,0.000033094515,0.00033600823,0.0007873162,0.2606773,0.00085584994,0.7316436,0.0048590363,0.000023672026],"about_ca_topic_score_codex":0.0024915815,"about_ca_topic_score_gemma":0.0024241772,"teacher_disagreement_score":0.004330215,"about_ca_system_score_codex":0.00062695646,"about_ca_system_score_gemma":0.0007319974,"threshold_uncertainty_score":0.014486015},"labels":[],"label_agreement":null},{"id":"W4388234122","doi":"10.1101/2023.10.30.564783","title":"Mining the neuroimaging literature","year":2023,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"","keywords":"Computer science; Workflow; Upload; Process (computing); Information retrieval; Task (project management); Information extraction; Data science; Biomedical text mining; Data mining; Text mining; World Wide Web; Database","score_opus":0.020110870889693648,"score_gpt":0.24594071357750957,"score_spread":0.22582984268781592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388234122","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09492111,0.09392494,0.30560562,0.026751151,0.0032508827,0.003908671,0.30360642,0.013323421,0.15470782],"genre_scores_gemma":[0.17047049,0.041466013,0.54659796,0.0026667903,0.0018937221,0.0029630777,0.2141499,0.002139161,0.017652912],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99521166,0.0012804123,0.00083084026,0.000930225,0.0015833164,0.0001635505],"domain_scores_gemma":[0.9752035,0.012308418,0.0026988634,0.0021626994,0.0068001,0.00082644826],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0060804584,0.00088839175,0.0008863672,0.0512669,0.0016090669,0.0040457617,0.0018266924,0.0010031144,0.011369827],"category_scores_gemma":[0.033482533,0.00045590306,0.0012370681,0.026605153,0.0010128577,0.0029962028,0.003533688,0.001045722,0.0056976327],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002208109,0.00012637838,0.013734693,0.0126307905,0.0004898143,0.0038415287,0.003036297,0.0016363104,0.010760737,0.028101053,0.1770854,0.74833614],"study_design_scores_gemma":[0.00007393129,0.00007178662,0.01707009,0.0061711287,0.0006948064,0.002518882,0.0028380465,0.00716459,0.011551088,0.0678634,0.88387954,0.0001027538],"about_ca_topic_score_codex":0.0047869845,"about_ca_topic_score_gemma":0.009229951,"teacher_disagreement_score":0.99391955,"about_ca_system_score_codex":0.0016733541,"about_ca_system_score_gemma":0.0069986903,"threshold_uncertainty_score":0.03803587},"labels":[],"label_agreement":null},{"id":"W4388253349","doi":"10.1093/oso/9780199638871.003.0006","title":"Internet tools for cell and developmental biologists","year":2004,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"The Internet; Computer science; Hyperlink; Simple (philosophy); Gastrulation; World Wide Web; Biology; Cell biology; Web page; Embryo","score_opus":0.04400597066745447,"score_gpt":0.25702656805508833,"score_spread":0.21302059738763385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388253349","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015929297,0.032132987,0.16646995,0.005922819,0.0025716966,0.0003161495,0.0035360868,0.014652192,0.7728052],"genre_scores_gemma":[0.008350486,0.044802517,0.20874985,0.004253745,0.0014709119,0.000712554,0.009605087,0.0040219952,0.7180329],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99962187,0.0000717058,0.000034670098,0.00005530964,0.00018920307,0.000027164306],"domain_scores_gemma":[0.9990754,0.00042215709,0.000029167282,0.00015655083,0.00022008989,0.00009654764],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.00072925206,0.0011435531,0.0005983927,0.0052697007,0.001044573,0.004650219,0.0017996596,0.0015702411,0.104004025],"category_scores_gemma":[0.001906514,0.00043854362,0.00059319485,0.0063195694,0.00084872346,0.006571467,0.002834757,0.0021217251,0.073336765],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013896609,0.00004097991,0.000122054276,0.00045118225,0.0000059998088,0.00026195147,0.00054447236,0.0003142181,0.002360165,0.13337757,0.3927126,0.46979496],"study_design_scores_gemma":[0.0000025009547,0.0000027738488,0.00007485411,0.00013055885,0.0000023137704,0.00019507235,0.000053063013,0.0001991362,0.00021644436,0.013291341,0.9858269,0.000004973012],"about_ca_topic_score_codex":0.0011886322,"about_ca_topic_score_gemma":0.0021980687,"teacher_disagreement_score":0.99534976,"about_ca_system_score_codex":0.0012508419,"about_ca_system_score_gemma":0.0009817801,"threshold_uncertainty_score":0.3479281},"labels":[],"label_agreement":null},{"id":"W4388362461","doi":"10.1016/j.ibneur.2023.08.1935","title":"A MULTI-OMICS APPROACH TO THE IDENTIFICATION OF BIOLOGICAL SIGNATURES IN PARKINSONISMS","year":2023,"lang":"en","type":"article","venue":"IBRO Neuroscience Reports","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Movement Disorders; Université de Montréal","funders":"","keywords":"Identification (biology); Computational biology; Omics; Computer science; Biology; Bioinformatics; Ecology","score_opus":0.055548953195949914,"score_gpt":0.31375806935413975,"score_spread":0.2582091161581898,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388362461","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.074016586,0.015874252,0.8755049,0.006729611,0.00036214385,0.0008243602,0.019588143,0.003308934,0.0037911],"genre_scores_gemma":[0.21604145,0.0075144726,0.75681585,0.0017637681,0.00039755975,0.0007999435,0.0150697,0.00024355449,0.001353715],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9980744,0.00066071434,0.000300335,0.0004645918,0.00039174652,0.000108196204],"domain_scores_gemma":[0.99761355,0.001224694,0.0002941871,0.00033057955,0.0003931147,0.0001437687],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004038296,0.0013678708,0.0018668381,0.010333037,0.00097432686,0.004015193,0.000877996,0.0009407301,0.0008486741],"category_scores_gemma":[0.0039121043,0.00040184212,0.0030085538,0.0062112915,0.0008758231,0.0021537766,0.0021025855,0.0014576148,0.00054158375],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013958763,0.0010300274,0.04313394,0.006397169,0.004252518,0.0033004866,0.0025109448,0.047034018,0.28729552,0.04329791,0.013633049,0.5467186],"study_design_scores_gemma":[0.00013833324,0.0005583902,0.099168815,0.0015311986,0.0025783456,0.0023074239,0.0036270313,0.39188257,0.050488994,0.36314702,0.08417225,0.0003995494],"about_ca_topic_score_codex":0.0032015033,"about_ca_topic_score_gemma":0.0028906022,"teacher_disagreement_score":0.010333037,"about_ca_system_score_codex":0.0012782892,"about_ca_system_score_gemma":0.0017576288,"threshold_uncertainty_score":0.021356761},"labels":[],"label_agreement":null},{"id":"W4388639867","doi":"10.1016/j.medj.2023.10.003","title":"The Medical Action Ontology: A tool for annotating and analyzing treatments and clinical management of human disease","year":2023,"lang":"en","type":"article","venue":"Med","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ottawa Hospital; Children's Hospital of Eastern Ontario; University of Ottawa","funders":"National Human Genome Research Institute; Medical Research Council; Memphis Research Consortium; National Institutes of Health; Evelyn Trust; Lily Foundation; National Institute for Health and Care Research; Department of Health and Social Care; NIHR Cambridge Biomedical Research Centre; Wellcome Trust; National Institute on Handicapped Research","keywords":"Ontology; Action (physics); Disease; Computer science; Knowledge management; Information retrieval; Data science; Medicine; Internal medicine; Epistemology","score_opus":0.05926305652662283,"score_gpt":0.42056887272198196,"score_spread":0.3613058161953591,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388639867","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0048696054,0.0034796605,0.7794023,0.0072906357,0.00082209124,0.0018720528,0.114600144,0.042842835,0.044820648],"genre_scores_gemma":[0.029824765,0.004548598,0.8478683,0.0021995436,0.0002490251,0.001721304,0.10323838,0.0041664797,0.0061835856],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99682933,0.00084125233,0.00078853645,0.0005333805,0.0008638713,0.00014358309],"domain_scores_gemma":[0.991708,0.004708195,0.00090151856,0.0012506173,0.00090454746,0.0005271229],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051352293,0.0016926143,0.0008294971,0.01132542,0.0022302873,0.003961068,0.0023747717,0.0024329454,0.016608996],"category_scores_gemma":[0.017033434,0.0011029914,0.0031169015,0.0060548466,0.0015853432,0.0069198045,0.006423661,0.0028820345,0.005613035],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040508126,0.000240911,0.008112808,0.008269966,0.00039640648,0.0019092449,0.0043311473,0.01403422,0.007224058,0.35679552,0.34969822,0.2485824],"study_design_scores_gemma":[0.00006540176,0.00003865791,0.0015149408,0.0013545784,0.000096742566,0.00074369356,0.0004910898,0.013182265,0.001993122,0.07539955,0.9050389,0.000081010476],"about_ca_topic_score_codex":0.015808009,"about_ca_topic_score_gemma":0.02199479,"teacher_disagreement_score":0.016608996,"about_ca_system_score_codex":0.003640967,"about_ca_system_score_gemma":0.007976051,"threshold_uncertainty_score":0.055562556},"labels":[],"label_agreement":null},{"id":"W4388760235","doi":"10.11647/obp.0335.04","title":"4. Models and Simulations","year":2023,"lang":"en","type":"book-chapter","venue":"Open Book Publishers","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Operations research; Perspective (graphical); Emulation; Action (physics); Domain (mathematical analysis); Management science; Data science; Artificial intelligence; Psychology; Engineering; Mathematics","score_opus":0.06479396447323817,"score_gpt":0.2984813054783758,"score_spread":0.23368734100513766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388760235","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008462532,0.011823371,0.48060045,0.032602906,0.003086388,0.0006716415,0.0052910475,0.0037626233,0.45369896],"genre_scores_gemma":[0.33481255,0.03334549,0.3488699,0.010573627,0.0032047604,0.0026159412,0.009067027,0.0033350193,0.25417578],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9964336,0.0012987028,0.00025163204,0.0005756156,0.0010855083,0.00035502357],"domain_scores_gemma":[0.995747,0.0020309195,0.0002987034,0.0007558663,0.00096826244,0.00019934366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027476298,0.0013663314,0.0012822413,0.0014480299,0.0020830906,0.009076284,0.0041258703,0.0059779524,0.04480164],"category_scores_gemma":[0.012666685,0.0008915751,0.002554149,0.0016494116,0.003421255,0.0076328856,0.0045790845,0.0049099163,0.015637886],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034320587,0.0000515927,0.00087199255,0.0004498774,0.000051302777,0.00013202567,0.00028943436,0.08267142,0.00028070871,0.86137515,0.028759295,0.025032897],"study_design_scores_gemma":[0.000037714566,0.000037490012,0.0002809687,0.0004305769,0.000027282447,0.00015299987,0.00022932235,0.08253435,0.00041084527,0.52933574,0.38647243,0.00005027638],"about_ca_topic_score_codex":0.014590905,"about_ca_topic_score_gemma":0.0049710753,"teacher_disagreement_score":0.04480164,"about_ca_system_score_codex":0.0037186402,"about_ca_system_score_gemma":0.0055724136,"threshold_uncertainty_score":0.14987648},"labels":[],"label_agreement":null},{"id":"W4388895410","doi":"10.2196/49041","title":"Extracting Clinical Information From Japanese Radiology Reports Using a 2-Stage Deep Learning Approach: Algorithm Development and Validation","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Deep learning; Computer science; Pipeline (software); Artificial intelligence; Information extraction; Stage (stratigraphy); Machine learning; Natural language processing; Information retrieval; Algorithm; Radiology; Medicine","score_opus":0.04765871925151699,"score_gpt":0.34557667999624603,"score_spread":0.29791796074472904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4388895410","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49521756,0.0022555108,0.47813797,0.0010471534,0.00019447824,0.0016130782,0.003458359,0.0153214615,0.0027545136],"genre_scores_gemma":[0.53014195,0.0006659082,0.4518381,0.00062548154,0.00006647313,0.0015364687,0.011436075,0.00019407342,0.0034954892],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99872977,0.00026898764,0.00019048848,0.00041147613,0.00025675035,0.0001425189],"domain_scores_gemma":[0.9970541,0.0014181992,0.00017378137,0.00028617756,0.0009574009,0.000110316876],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0032407208,0.0022329404,0.0009260861,0.0019075169,0.0006079641,0.0009776428,0.0022629846,0.002176989,0.0018207827],"category_scores_gemma":[0.006066877,0.0007234714,0.0012666502,0.0011987854,0.00041817056,0.0013654876,0.001477789,0.0016373621,0.00089937577],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007639481,0.001456946,0.020931108,0.000581155,0.00040097264,0.0004290791,0.00021658845,0.28116843,0.022162465,0.0009185527,0.010161195,0.6608096],"study_design_scores_gemma":[0.00006464304,0.00014759647,0.0020884897,0.000021384185,0.000050450028,0.00007103266,0.000030153025,0.9874511,0.00889472,0.00050525524,0.00065895414,0.000016109241],"about_ca_topic_score_codex":0.016846841,"about_ca_topic_score_gemma":0.018139254,"teacher_disagreement_score":0.016846841,"about_ca_system_score_codex":0.0017281017,"about_ca_system_score_gemma":0.0033065954,"threshold_uncertainty_score":0.033497572},"labels":[],"label_agreement":null},{"id":"W4389083881","doi":"10.2196/44639","title":"Patient Information Summarization in Clinical Settings: Scoping Review","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Automatic summarization; Computer science; MEDLINE; Medicine; Data science; Information retrieval","score_opus":0.023312361755187333,"score_gpt":0.3673106231275995,"score_spread":0.3439982613724122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389083881","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017908101,0.9849819,0.0043655,0.0032307059,0.0006399201,0.002330189,0.00056812423,0.00006425691,0.0020285926],"genre_scores_gemma":[0.023697432,0.9541998,0.014828353,0.0014981369,0.0004792376,0.004271278,0.0007690104,0.00003772167,0.00021900123],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.90728605,0.04493636,0.032577414,0.0031519346,0.011167121,0.0008811211],"domain_scores_gemma":[0.5351257,0.38723487,0.03655898,0.009280383,0.030564617,0.0012354334],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.08890187,0.002365891,0.0058619664,0.039402388,0.0022024016,0.0089740185,0.0041015046,0.004965403,0.0037781294],"category_scores_gemma":[0.33178666,0.0017971608,0.0069148485,0.03563578,0.003294933,0.0093865665,0.004921798,0.0030548484,0.0007960712],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001853181,0.00007729076,0.0013655556,0.68850523,0.0024386346,0.00019997972,0.002524065,0.0007173491,0.0003181327,0.002827036,0.0060857325,0.29475576],"study_design_scores_gemma":[0.000053190863,0.00015001548,0.0013862492,0.9516615,0.006384459,0.00028322625,0.0014830742,0.00037399665,0.00048004222,0.0020113064,0.03568149,0.000051363448],"about_ca_topic_score_codex":0.005440771,"about_ca_topic_score_gemma":0.008527502,"teacher_disagreement_score":0.08890187,"about_ca_system_score_codex":0.008005046,"about_ca_system_score_gemma":0.03486011,"threshold_uncertainty_score":0.47016364},"labels":[],"label_agreement":null},{"id":"W4389464289","doi":"10.1212/wnl.92.15_supplement.p5.8-024","title":"What Bothers Parkinson Patients: Clinical Curation of their Verbatim Reports (P5.8-024)","year":2019,"lang":"en","type":"article","venue":"Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network","funders":"","keywords":"Medicine","score_opus":0.01497914585075859,"score_gpt":0.2823230816531881,"score_spread":0.2673439358024295,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389464289","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09981523,0.008162744,0.025840187,0.042890646,0.0030147394,0.0022579248,0.76281834,0.009596606,0.04560353],"genre_scores_gemma":[0.29340518,0.0057549556,0.09037418,0.010860093,0.001309261,0.004038107,0.5617168,0.004239856,0.028301638],"study_design_codex":"not_applicable","study_design_gemma":"qualitative","domain_scores_codex":[0.996813,0.0009714803,0.0007133475,0.00050340145,0.0008018424,0.00019698137],"domain_scores_gemma":[0.97674245,0.012934798,0.0025847487,0.0012071071,0.005641283,0.00088958966],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003777606,0.00051119726,0.0006565164,0.00554898,0.0009901955,0.0022055348,0.0007653667,0.0012033592,0.022597004],"category_scores_gemma":[0.048401333,0.0002868147,0.0008233579,0.0028448529,0.00037976488,0.001681305,0.0026016876,0.000845082,0.010899104],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004268444,0.00009994817,0.031735703,0.004618776,0.00016292604,0.0011176675,0.002896094,0.000197367,0.0028880942,0.0012828759,0.81856894,0.13600475],"study_design_scores_gemma":[0.00026652354,0.00011569599,0.072378986,0.0039649224,0.00043067502,0.0030224659,0.0037668827,0.0023081326,0.006518519,0.005709088,0.90138876,0.00012920935],"about_ca_topic_score_codex":0.009064159,"about_ca_topic_score_gemma":0.013970417,"teacher_disagreement_score":0.022597004,"about_ca_system_score_codex":0.0011050628,"about_ca_system_score_gemma":0.003045005,"threshold_uncertainty_score":0.075594544},"labels":[],"label_agreement":null},{"id":"W4389543603","doi":"10.1109/ichi57859.2023.00134","title":"Investigation into Scaling-Up the SOAP Problem-Oriented Medical Record into a Clinical Case Study","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"SOAP; Clinical Practice; Computer science; Health care; Medical education; Scale (ratio); Medical record; Data science; Medicine; World Wide Web; Nursing","score_opus":0.05971284833956395,"score_gpt":0.3808122086411614,"score_spread":0.3210993603015974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4389543603","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4721556,0.00092451187,0.442633,0.018810263,0.0007954323,0.011068357,0.0024372095,0.028622065,0.022553574],"genre_scores_gemma":[0.2371379,0.00034018836,0.75601715,0.000792446,0.00004576414,0.0010119858,0.0017346917,0.00070768327,0.0022122054],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.98472035,0.007576782,0.001935004,0.0017560496,0.0034937658,0.00051802624],"domain_scores_gemma":[0.89400935,0.07086912,0.0028392768,0.017315406,0.011805345,0.003161506],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026541328,0.0006142345,0.000501635,0.0018507353,0.0013536749,0.005451565,0.0033384203,0.0014114133,0.003290531],"category_scores_gemma":[0.08375863,0.0006532846,0.0012399737,0.0016328224,0.0019001258,0.00756294,0.004381519,0.00294022,0.001052524],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017732277,0.0050050183,0.054674853,0.0042425846,0.0003646502,0.006762196,0.044039037,0.02383296,0.066033944,0.042310216,0.02721686,0.7237444],"study_design_scores_gemma":[0.0014463331,0.005346906,0.04086361,0.0029147714,0.0006804164,0.007816488,0.040861096,0.3681726,0.09215824,0.044684574,0.39438403,0.0006709442],"about_ca_topic_score_codex":0.006543913,"about_ca_topic_score_gemma":0.008441243,"teacher_disagreement_score":0.026541328,"about_ca_system_score_codex":0.002413997,"about_ca_system_score_gemma":0.0044773594,"threshold_uncertainty_score":0.1403656},"labels":[],"label_agreement":null},{"id":"W4390114278","doi":"10.2196/49301","title":"The Necessity of Interoperability to Uncover the Full Potential of Digital Health Devices","year":2023,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Bundesministerium für Bildung und Forschung","keywords":"SNOMED CT; Interoperability; Systematized Nomenclature of Medicine; Identifier; Medicine; Computer science; Semantic interoperability; Health care; Concordance; Medical physics; Terminology; World Wide Web","score_opus":0.014762477487795498,"score_gpt":0.31552672823001554,"score_spread":0.30076425074222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390114278","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03197107,0.0075459136,0.8582528,0.042992905,0.0010649743,0.0011618999,0.0027274108,0.0024247083,0.05185826],"genre_scores_gemma":[0.26096547,0.0059894673,0.70914054,0.008588868,0.0010850086,0.0014797082,0.0060610287,0.00095199666,0.005737944],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9433484,0.028126014,0.0075546103,0.0053248527,0.014202721,0.0014434659],"domain_scores_gemma":[0.8303028,0.085343935,0.005824141,0.06140285,0.015714927,0.0014113577],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.064050145,0.0013558988,0.0016873209,0.008623693,0.002171409,0.012760401,0.0041295006,0.0032067816,0.004346883],"category_scores_gemma":[0.1196714,0.0012672523,0.0022227576,0.007134887,0.0073776753,0.036491577,0.013682997,0.0072464547,0.002235296],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021537834,0.00025816236,0.016264847,0.002526294,0.00050588965,0.00084016274,0.00948647,0.0027960124,0.0059479247,0.4798412,0.012932031,0.46838552],"study_design_scores_gemma":[0.00003888456,0.00016888733,0.008906869,0.0035247968,0.0003659189,0.0016038566,0.005217202,0.00825475,0.0045776935,0.7183547,0.24882671,0.00015971261],"about_ca_topic_score_codex":0.002851692,"about_ca_topic_score_gemma":0.0017594896,"teacher_disagreement_score":0.064050145,"about_ca_system_score_codex":0.002453873,"about_ca_system_score_gemma":0.005367513,"threshold_uncertainty_score":0.3387336},"labels":[],"label_agreement":null},{"id":"W4390712701","doi":"10.2196/49607","title":"Impact of Translation on Biomedical Information Extraction: Experiment on Real-Life Clinical Notes","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Translation (biology); Information extraction; Information retrieval; Natural language processing; Data science; Data mining; Artificial intelligence; Biology","score_opus":0.0523874128518099,"score_gpt":0.43424987324319636,"score_spread":0.38186246039138644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390712701","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.88765204,0.0062712724,0.06974031,0.001618626,0.0007889576,0.0011966194,0.0067486954,0.017692076,0.008291381],"genre_scores_gemma":[0.81690484,0.0019624792,0.14739911,0.0010092021,0.0002648852,0.0005288817,0.026415365,0.00094735133,0.0045678383],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99334866,0.003645581,0.0007073984,0.0012388198,0.00080418773,0.00025532753],"domain_scores_gemma":[0.9821873,0.0128123835,0.0005184825,0.0015863386,0.0025680237,0.00032750217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008510824,0.0024124437,0.0011029021,0.0022209424,0.0011481149,0.001311127,0.0011774915,0.0016868392,0.0029401612],"category_scores_gemma":[0.025862928,0.00035454214,0.001218499,0.0023624732,0.0008764167,0.0021356978,0.0015065459,0.0010074246,0.0024120766],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0056966064,0.0042944606,0.040401172,0.004297839,0.001706871,0.0021510446,0.0018699886,0.062798865,0.05430004,0.0013135385,0.03682233,0.78434724],"study_design_scores_gemma":[0.002514373,0.007915593,0.124628805,0.0008854861,0.002674777,0.0055961907,0.0037171827,0.51188934,0.27905768,0.0055471603,0.054966774,0.00060660276],"about_ca_topic_score_codex":0.016631238,"about_ca_topic_score_gemma":0.014414757,"teacher_disagreement_score":0.016631238,"about_ca_system_score_codex":0.0012659837,"about_ca_system_score_gemma":0.0014680638,"threshold_uncertainty_score":0.04501009},"labels":[],"label_agreement":null},{"id":"W4390865012","doi":"","title":"What can you do next? Choice of output and reuse of your transcription","year":2023,"lang":"en","type":"preprint","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Reuse; Transcription (linguistics); Computer science; Engineering; Waste management","score_opus":0.048026670050582344,"score_gpt":0.2778761987056975,"score_spread":0.22984952865511518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390865012","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056802295,0.0010908276,0.789732,0.016736913,0.0033046491,0.00036484594,0.017608924,0.033902653,0.08045683],"genre_scores_gemma":[0.40551558,0.0016652226,0.41328132,0.0019238857,0.001064793,0.00071339816,0.027076975,0.044967487,0.1037912],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9903526,0.0035877305,0.0007710405,0.0019516482,0.0024813674,0.0008556],"domain_scores_gemma":[0.9629549,0.014320841,0.0007931084,0.011380218,0.009204624,0.0013463787],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009977587,0.0009793018,0.0012498517,0.0014892465,0.0016173996,0.006055702,0.0016084603,0.0016585112,0.036464736],"category_scores_gemma":[0.051683594,0.00069118146,0.0020663843,0.0021101767,0.0017583006,0.007502655,0.0037821017,0.0025675176,0.038573697],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0031954092,0.00031498497,0.017696809,0.002078648,0.0002925401,0.0018913577,0.01935166,0.005005917,0.053125158,0.08558786,0.18625017,0.62520945],"study_design_scores_gemma":[0.00018985612,0.00030916304,0.00873763,0.001393038,0.0003890007,0.0015905313,0.008223377,0.02862431,0.11957615,0.12134603,0.7092052,0.00041576094],"about_ca_topic_score_codex":0.00321605,"about_ca_topic_score_gemma":0.0024118721,"teacher_disagreement_score":0.036464736,"about_ca_system_score_codex":0.0013818023,"about_ca_system_score_gemma":0.0029850816,"threshold_uncertainty_score":0.12198663},"labels":[],"label_agreement":null},{"id":"W4390888571","doi":"10.17504/protocols.io.5qpvo3wodv4o/v1","title":"Protocol for Systematic Analysis Data Elements v1","year":2023,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Interoperability; Computer science; Harmonization; Standardization; Protocol (science); Credibility; Data quality; Best practice; Data science; World Wide Web; Engineering; Medicine","score_opus":0.18749455231755605,"score_gpt":0.4458937561904043,"score_spread":0.2583992038728482,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4390888571","genre_codex":"protocol","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":"protocol","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011844094,0.0009339707,0.13549669,0.006190579,0.004383141,0.51278925,0.30080324,0.009372721,0.028845983],"genre_scores_gemma":[0.001814858,0.0003503068,0.09226735,0.0014335407,0.00019753413,0.8718555,0.023527661,0.001036525,0.0075167515],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.84216475,0.0789077,0.048478067,0.010577638,0.015287337,0.0045843865],"domain_scores_gemma":[0.58425105,0.20320484,0.016382286,0.11360027,0.0755925,0.0069690174],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.17327954,0.0039578523,0.0059977453,0.020650059,0.006349761,0.012077738,0.005903981,0.008680863,0.3846772],"category_scores_gemma":[0.3076425,0.007074319,0.00578793,0.017676435,0.0077917646,0.0070439205,0.012116435,0.014468446,0.12825039],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029646815,0.0003853041,0.0013551915,0.05602036,0.0006176011,0.0007742538,0.007048045,0.0013990875,0.0033813051,0.0753077,0.7444064,0.1063401],"study_design_scores_gemma":[0.0027395077,0.00019406121,0.0015671159,0.02098057,0.00020432798,0.00019808867,0.0018264254,0.0010501075,0.0022860472,0.054112814,0.9145115,0.0003294554],"about_ca_topic_score_codex":0.0058845324,"about_ca_topic_score_gemma":0.006426303,"teacher_disagreement_score":0.8267205,"about_ca_system_score_codex":0.009845565,"about_ca_system_score_gemma":0.08037208,"threshold_uncertainty_score":0.91640073},"labels":[],"label_agreement":null},{"id":"W4391070609","doi":"10.1016/j.jbi.2024.104588","title":"Semantics-enabled biomedical literature analytics","year":2024,"lang":"en","type":"editorial","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University; University of Toronto","funders":"U.S. National Library of Medicine; National Institutes of Health","keywords":"Computer science; Analytics; Semantics (computer science); Data science; Programming language","score_opus":0.007967348372261453,"score_gpt":0.28493941550784746,"score_spread":0.27697206713558603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391070609","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00013788743,0.011521153,0.0048810714,0.0876063,0.89196813,0.00004061201,0.00056639564,0.0005683406,0.0027100404],"genre_scores_gemma":[0.0025537377,0.014174097,0.0044048657,0.02112786,0.94478,0.000058564267,0.0005592814,0.00023427387,0.012107294],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9930397,0.0016391085,0.0009765019,0.0006616899,0.0034434174,0.00023954938],"domain_scores_gemma":[0.94831634,0.031089464,0.0015343702,0.0012143565,0.01472136,0.0031241267],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.012538124,0.0021108375,0.002683418,0.010075828,0.002139575,0.013358208,0.0030554845,0.007529995,0.00836453],"category_scores_gemma":[0.03646446,0.0014954103,0.002264459,0.0036503216,0.0034356753,0.009076821,0.003200978,0.013263818,0.004549611],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000053915213,0.0000151145005,0.000038674712,0.0003335147,0.00005969817,0.0000955347,0.000022625012,0.00012522878,0.00013755605,0.0022198518,0.9801128,0.016785389],"study_design_scores_gemma":[0.0000908186,0.000019910398,0.00020089334,0.0006824822,0.00013025028,0.00022907693,0.000050578878,0.0019267687,0.00046215934,0.017189875,0.9789781,0.000039252984],"about_ca_topic_score_codex":0.0019420091,"about_ca_topic_score_gemma":0.0052625346,"teacher_disagreement_score":0.98746186,"about_ca_system_score_codex":0.0027309721,"about_ca_system_score_gemma":0.004272517,"threshold_uncertainty_score":0.06630874},"labels":[],"label_agreement":null},{"id":"W4391093090","doi":"10.1109/bigdata59044.2023.10386643","title":"Text mining using clinical terms in electronic records of annual falls of patients in home community care","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Medicine","score_opus":0.0346535202987868,"score_gpt":0.3469227388240934,"score_spread":0.3122692185253066,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391093090","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8380242,0.0030263932,0.020511752,0.0013079874,0.00029970065,0.001078173,0.12664114,0.0020473013,0.0070633707],"genre_scores_gemma":[0.73483115,0.001458209,0.09782925,0.00051887176,0.00023314296,0.0011988444,0.16069306,0.0001904108,0.003047212],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9979765,0.00031571373,0.00052233797,0.0005857913,0.00044933546,0.000150261],"domain_scores_gemma":[0.99550086,0.0025192196,0.0008537741,0.00025870075,0.00068445114,0.00018297351],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010484303,0.0006492476,0.0005241555,0.009596809,0.0007462999,0.0012648562,0.0006175231,0.00095534686,0.0017312239],"category_scores_gemma":[0.008301516,0.00019556242,0.00083727826,0.007677451,0.00044461913,0.0009977734,0.0012106403,0.0006388044,0.0015156601],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019319097,0.0009022926,0.40914017,0.006862428,0.00048068468,0.009039331,0.007650359,0.007033358,0.044661954,0.0029943027,0.070791684,0.43851158],"study_design_scores_gemma":[0.00028515293,0.00082051277,0.7138642,0.0015458785,0.0006499822,0.010156428,0.011427229,0.07422368,0.028995087,0.008220942,0.14953241,0.0002784961],"about_ca_topic_score_codex":0.010524621,"about_ca_topic_score_gemma":0.020998871,"teacher_disagreement_score":0.010524621,"about_ca_system_score_codex":0.0008540884,"about_ca_system_score_gemma":0.00187189,"threshold_uncertainty_score":0.020926714},"labels":[],"label_agreement":null},{"id":"W4391224465","doi":"10.3233/shti230943","title":"An Ontology-Based Architecture to Support Language Variants of Model-Driven Electronic Health Records","year":2024,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Ontology; Natural language processing; Mandarin Chinese; German; Focus (optics); Portuguese; Artificial intelligence; Linguistics","score_opus":0.025899301811110675,"score_gpt":0.3813000614387498,"score_spread":0.3554007596276391,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391224465","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006744418,0.00014586945,0.9797717,0.0010764323,0.00009367563,0.0003583334,0.00034080862,0.008173094,0.0032958027],"genre_scores_gemma":[0.07338171,0.00038243472,0.91705173,0.0005976398,0.000049956463,0.00034782218,0.0028844601,0.0012848359,0.004019332],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99620444,0.0010016004,0.0007716979,0.0006025758,0.0011181143,0.00030166365],"domain_scores_gemma":[0.99524045,0.0014433624,0.00033332861,0.0015450412,0.0010127957,0.00042490664],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0091806585,0.000644448,0.00079933525,0.002438711,0.0015453913,0.005392531,0.0033174169,0.00201655,0.0022062275],"category_scores_gemma":[0.011823679,0.0011275391,0.0024735632,0.0023271479,0.00163027,0.00998423,0.0053997478,0.0034548761,0.0012813581],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00059457717,0.0009580034,0.010550331,0.0010005351,0.0005593481,0.0028766468,0.011445801,0.05043725,0.028943093,0.48799574,0.031669587,0.37296912],"study_design_scores_gemma":[0.00017119467,0.00017901891,0.0022041867,0.00070878305,0.00046857152,0.001300987,0.0013432399,0.40893927,0.020968383,0.1816824,0.38168508,0.00034893482],"about_ca_topic_score_codex":0.017408112,"about_ca_topic_score_gemma":0.017653663,"teacher_disagreement_score":0.017408112,"about_ca_system_score_codex":0.0027184996,"about_ca_system_score_gemma":0.005251635,"threshold_uncertainty_score":0.048552573},"labels":[],"label_agreement":null},{"id":"W4391348681","doi":"10.2196/53516","title":"Using a Natural Language Processing Approach to Support Rapid Knowledge Acquisition","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Terminology; Documentation; Computer science; Cloud computing; Phenome; Health care; Clinical decision support system; Knowledge management; Electronic health record; Health informatics; Data science; Artificial intelligence; Natural language processing; Decision support system; Medicine; Pathology; Linguistics; Public health","score_opus":0.026386628124518263,"score_gpt":0.35273507264230247,"score_spread":0.3263484445177842,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391348681","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066427686,0.00027178018,0.9733925,0.0019446191,0.000097425225,0.0011959675,0.0022722178,0.008194413,0.005988287],"genre_scores_gemma":[0.02873971,0.0002685554,0.96575737,0.0004958974,0.000048270154,0.0006776914,0.0025989471,0.00022883256,0.0011846609],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9945018,0.0024025412,0.00065980473,0.0010609201,0.0012479938,0.00012693016],"domain_scores_gemma":[0.96933144,0.023119094,0.0012423827,0.0028349063,0.0031913402,0.00028081157],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0074029616,0.001458802,0.0009153355,0.005285725,0.0014574502,0.0046985727,0.0023636995,0.0012359251,0.0067157904],"category_scores_gemma":[0.029540531,0.00082975096,0.0017854271,0.0033839103,0.0016940879,0.006304804,0.004237436,0.0026763654,0.0038924187],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003528811,0.00095414114,0.0033032512,0.003504682,0.00032596663,0.0019788542,0.0060641672,0.018385962,0.042192403,0.088688254,0.03646297,0.7977865],"study_design_scores_gemma":[0.00030214209,0.00045712487,0.0025938374,0.001223563,0.00033253338,0.0013867812,0.003362418,0.31388327,0.047658246,0.37914538,0.24932769,0.00032695435],"about_ca_topic_score_codex":0.006221298,"about_ca_topic_score_gemma":0.012503316,"teacher_disagreement_score":0.0074029616,"about_ca_system_score_codex":0.0016431584,"about_ca_system_score_gemma":0.0043329354,"threshold_uncertainty_score":0.039151073},"labels":[],"label_agreement":null},{"id":"W4391733408","doi":"10.1016/j.knosys.2024.111493","title":"Plausible reasoning over large health datasets: A novel approach to data analytics leveraging semantics","year":2024,"lang":"en","type":"article","venue":"Knowledge-Based Systems","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Computer science; Semantics (computer science); Knowledge representation and reasoning; Ontology; Question answering; Information retrieval; Semantic Web; Analytics; Natural language processing; Artificial intelligence; Data science","score_opus":0.0843420894697554,"score_gpt":0.35170762492839,"score_spread":0.26736553545863456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391733408","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062721525,0.0009842658,0.98622185,0.0019833348,0.00007716273,0.0002465727,0.0020357464,0.0011317286,0.0010471572],"genre_scores_gemma":[0.1563643,0.0010555704,0.8348838,0.0007946745,0.0002883926,0.00031667357,0.0052819964,0.0001874619,0.00082704035],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98890555,0.0031099517,0.0014720918,0.00224772,0.0039478014,0.00031691298],"domain_scores_gemma":[0.95908433,0.031107878,0.0021893785,0.0051332163,0.0017833239,0.0007018393],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009859375,0.0016393415,0.002708905,0.0081717195,0.0019814756,0.009027145,0.005192848,0.0029487102,0.00318242],"category_scores_gemma":[0.053291045,0.001470887,0.0061620143,0.008388418,0.0027071508,0.013786752,0.008843697,0.0046955524,0.0007756393],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090640975,0.00079096545,0.011773409,0.0027842887,0.0019872452,0.0037409998,0.0025862663,0.13658822,0.0071957298,0.41316444,0.018806573,0.3996755],"study_design_scores_gemma":[0.00008489593,0.00006701674,0.00068384956,0.00017767827,0.00031967636,0.0005739771,0.00033990416,0.44028842,0.0018609312,0.54558754,0.009961636,0.00005448548],"about_ca_topic_score_codex":0.0048297676,"about_ca_topic_score_gemma":0.00920482,"teacher_disagreement_score":0.009859375,"about_ca_system_score_codex":0.0015941664,"about_ca_system_score_gemma":0.0039685494,"threshold_uncertainty_score":0.052141964},"labels":[],"label_agreement":null},{"id":"W4391755833","doi":"10.1016/j.isci.2024.109212","title":"Distance-weighted Sinkhorn loss for Alzheimer’s disease classification","year":2024,"lang":"en","type":"article","venue":"iScience","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Janssen Alzheimer Immunotherapy Research And Development; Canadian Institutes of Health Research; National Institute of Biomedical Imaging and Bioengineering; National Institutes of Health; Genentech; IXICO; U.S. National Library of Medicine; H. Lundbeck A/S; Servier; Eisai; Northern California Institute for Research and Education; University of Southern California; Pfizer; Biogen; BioClinica; Meso Scale Diagnostics; Eli Lilly and Company; U.S. Department of Defense; Alzheimer's Disease Neuroimaging Initiative; Novartis Pharmaceuticals Corporation; Bristol-Myers Squibb; National Institute on Aging; Alzheimer's Association; National Science Foundation","keywords":"Alzheimer's disease; Disease; Computer science; Medicine; Internal medicine","score_opus":0.043552517419552614,"score_gpt":0.32656143857810943,"score_spread":0.2830089211585568,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4391755833","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2552431,0.0042010504,0.73060536,0.0019651563,0.00026598314,0.00017276684,0.0010065946,0.0013951693,0.005144854],"genre_scores_gemma":[0.893933,0.0008972348,0.097055465,0.0004891444,0.00020712831,0.00014795679,0.0022402694,0.00013203309,0.0048977314],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99857867,0.00046873163,0.00012652043,0.00026255051,0.0004619873,0.000101558246],"domain_scores_gemma":[0.99679846,0.0015689452,0.00035550262,0.00042694822,0.00067219394,0.00017796732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005355932,0.0010642887,0.001101085,0.001482815,0.0005431131,0.0011311722,0.00132555,0.0014831716,0.0011515042],"category_scores_gemma":[0.0076070316,0.00021058708,0.0006131159,0.0011809382,0.0011748457,0.0021871785,0.0015000699,0.001467488,0.000510729],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014351447,0.00073661155,0.029422779,0.00031777163,0.00032295604,0.0004591348,0.00018327284,0.50789356,0.007129836,0.018414881,0.022236533,0.41144753],"study_design_scores_gemma":[0.0000299233,0.00019262545,0.002772098,0.000033375072,0.00003129783,0.00019207971,0.000037127633,0.9739532,0.0026921127,0.018402714,0.0016453336,0.000018187291],"about_ca_topic_score_codex":0.0014747516,"about_ca_topic_score_gemma":0.0013734788,"teacher_disagreement_score":0.005355932,"about_ca_system_score_codex":0.0010088077,"about_ca_system_score_gemma":0.0011814265,"threshold_uncertainty_score":0.02832526},"labels":[],"label_agreement":null},{"id":"W4392068448","doi":"10.1016/j.jbi.2024.104614","title":"Improving the interoperability of drugs terminologies: Infusing local standardization with an international perspective","year":2024,"lang":"en","type":"article","venue":"Journal of Biomedical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Centre Hospitalier de l’Université de Montréal; Université de Montréal; Centre Intégré Universitaire de Santé et de Services Sociaux du Centre-Sud-de-l'Île-de-Montréal","funders":"Institut de Valorisation des Données; Canada First Research Excellence Fund","keywords":"SNOMED CT; Interoperability; Computer science; Standardization; Ontology; Semantic interoperability; Service (business); Code (set theory); Terminology; Systematized Nomenclature of Medicine; Information retrieval; World Wide Web; Software engineering; Programming language; Business; Linguistics","score_opus":0.011151556754645079,"score_gpt":0.2940526505757637,"score_spread":0.2829010938211186,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392068448","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06414608,0.009649746,0.6516802,0.056248747,0.0017139551,0.0012174388,0.0030012995,0.0036398007,0.20870277],"genre_scores_gemma":[0.36767718,0.007888232,0.58190316,0.01111225,0.000651781,0.0007000189,0.008563807,0.0027886962,0.018714827],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.93263555,0.030425116,0.0076889056,0.005391444,0.020647738,0.0032112955],"domain_scores_gemma":[0.8732483,0.033314653,0.006950631,0.045217935,0.03825342,0.0030151147],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.071632616,0.0010226449,0.0012066753,0.009453688,0.003965802,0.014959623,0.00317143,0.0021379741,0.0033919944],"category_scores_gemma":[0.08637467,0.00083662703,0.0017897545,0.013794716,0.0071163424,0.021658951,0.012248894,0.0053291763,0.0015240206],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018903473,0.00023143899,0.019939188,0.0013180698,0.00025679125,0.0004450639,0.017406126,0.0047638887,0.006993119,0.5594966,0.02763074,0.3613299],"study_design_scores_gemma":[0.00004788976,0.00013629168,0.012408186,0.0033382024,0.0003874963,0.00053051184,0.014632635,0.008530706,0.009957629,0.096251056,0.8536151,0.00016428206],"about_ca_topic_score_codex":0.21006462,"about_ca_topic_score_gemma":0.18086796,"teacher_disagreement_score":0.21006462,"about_ca_system_score_codex":0.023599617,"about_ca_system_score_gemma":0.06817497,"threshold_uncertainty_score":0.41768384},"labels":[],"label_agreement":null},{"id":"W4392147016","doi":"10.1186/s13326-024-00302-5","title":"Enriching the FIDEO ontology with food-drug interactions from online knowledge sources","year":2024,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"Horizon 2020 Framework Programme; Agence Nationale de la Recherche","keywords":"DrugBank; Ontology; Computer science; Data science; Information retrieval; Drug; Medicine","score_opus":0.015153185082832654,"score_gpt":0.2918638677920983,"score_spread":0.27671068270926563,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392147016","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0757854,0.008327216,0.74858195,0.014203126,0.0011779417,0.0019683635,0.086387426,0.009335239,0.054233372],"genre_scores_gemma":[0.14767137,0.008953147,0.7195367,0.00236286,0.0003510323,0.00089620607,0.11086534,0.0010356244,0.0083276965],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9974462,0.0005785492,0.00053290394,0.0004415467,0.0008133781,0.00018739248],"domain_scores_gemma":[0.9950513,0.0025929206,0.0004638669,0.00071260217,0.00086671003,0.00031245183],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034017928,0.0009269106,0.0009384374,0.015407738,0.0016564976,0.004243952,0.0011775162,0.0012290312,0.003287475],"category_scores_gemma":[0.009702555,0.0005445726,0.002362328,0.0099583585,0.0012395728,0.007817479,0.004526932,0.0015905686,0.0012117594],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005359713,0.0006717684,0.02199957,0.007948578,0.00080044474,0.0065985606,0.010065452,0.010698751,0.019186862,0.27131513,0.06600646,0.5841724],"study_design_scores_gemma":[0.000069235124,0.000067656,0.011361546,0.0027972613,0.0004991665,0.0025768965,0.003561129,0.022763943,0.006676106,0.09699116,0.8524664,0.0001694782],"about_ca_topic_score_codex":0.016025815,"about_ca_topic_score_gemma":0.024970172,"teacher_disagreement_score":0.016025815,"about_ca_system_score_codex":0.0028318588,"about_ca_system_score_gemma":0.0060369405,"threshold_uncertainty_score":0.03186506},"labels":[],"label_agreement":null},{"id":"W4392190584","doi":"10.1186/s12859-024-05693-x","title":"GPAD: a natural language processing-based application to extract the gene-disease association discovery information from OMIM","year":2024,"lang":"en","type":"article","venue":"BMC Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Alberta Children's Hospital; University of Calgary","funders":"National Human Genome Research Institute; Canadian Institutes of Health Research; Genome British Columbia; Compute Canada; Genome Canada","keywords":"Computational biology; Disease; DNA microarray; Association (psychology); Gene; Natural (archaeology); Computer science; Biology; Data science; Bioinformatics; Genetics; Medicine; Gene expression; Psychology; Pathology","score_opus":0.006096063725336593,"score_gpt":0.25596481010590466,"score_spread":0.24986874638056808,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392190584","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008501188,0.00080701034,0.33503368,0.0017597618,0.00031012893,0.0012332342,0.18660247,0.4606575,0.005095144],"genre_scores_gemma":[0.052930053,0.0011182443,0.7093519,0.00215583,0.00025988615,0.0026542179,0.21210304,0.015509825,0.003917134],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977245,0.000496647,0.00050671474,0.00065258023,0.00053772266,0.00008190698],"domain_scores_gemma":[0.9919659,0.0058250264,0.00073016644,0.0006449177,0.0006055308,0.00022841584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030232985,0.0031149432,0.0009029394,0.0055722627,0.00079701014,0.0027862473,0.0020116398,0.0014429745,0.021368464],"category_scores_gemma":[0.013975318,0.0011358038,0.0028107397,0.0028502906,0.00088989135,0.002998192,0.004250119,0.0023492645,0.013663343],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001708929,0.00047622717,0.013173853,0.0116891395,0.0011089728,0.0066949846,0.004165384,0.011830695,0.050425306,0.020551013,0.4619996,0.4161759],"study_design_scores_gemma":[0.0010276454,0.0005609373,0.015057307,0.0011204937,0.00046736666,0.005874367,0.0013205678,0.19557664,0.055036128,0.074202046,0.6491471,0.0006094698],"about_ca_topic_score_codex":0.0031216238,"about_ca_topic_score_gemma":0.0036623527,"teacher_disagreement_score":0.021368464,"about_ca_system_score_codex":0.0010564367,"about_ca_system_score_gemma":0.0022603816,"threshold_uncertainty_score":0.071484685},"labels":[],"label_agreement":null},{"id":"W4392238058","doi":"10.7717/peerj-cs.1888","title":"exKidneyBERT: a language model for kidney transplant pathology reports and the crucial role of extended vocabularies","year":2024,"lang":"en","type":"article","venue":"PeerJ Computer Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Vocabulary; Natural language processing; Artificial intelligence; Information extraction; Information retrieval; Language model; Domain (mathematical analysis); Field (mathematics); Pathology; Medicine; Linguistics","score_opus":0.007218420904766878,"score_gpt":0.25778343054216063,"score_spread":0.25056500963739375,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392238058","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2736341,0.003383613,0.67048633,0.0026888018,0.0007720491,0.00086041284,0.013238916,0.025323668,0.009612148],"genre_scores_gemma":[0.7788006,0.00088986306,0.19082677,0.0006839468,0.0001405636,0.00063773897,0.016122868,0.00051251834,0.011385088],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99945325,0.00016345843,0.000058996025,0.00019404154,0.000078408106,0.00005192383],"domain_scores_gemma":[0.9983858,0.0010513072,0.00012718623,0.000115204806,0.00025753552,0.00006291409],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017081993,0.0011552948,0.0005559561,0.0013590086,0.0003840123,0.00129988,0.0020167285,0.0008607477,0.0035513584],"category_scores_gemma":[0.0042662304,0.0005255826,0.0010711679,0.00060284836,0.00034799497,0.002736073,0.0010613853,0.0017888637,0.0017033662],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001680921,0.0004345442,0.009794366,0.00078656676,0.00042102288,0.0009963886,0.00083023205,0.38017967,0.020586,0.010846204,0.03455522,0.5388888],"study_design_scores_gemma":[0.000054726883,0.000088098284,0.0010195157,0.000053877957,0.00007522027,0.00015417293,0.0000709555,0.9854299,0.0050222725,0.003194845,0.004796907,0.000039539416],"about_ca_topic_score_codex":0.026720487,"about_ca_topic_score_gemma":0.026391469,"teacher_disagreement_score":0.026720487,"about_ca_system_score_codex":0.0019542268,"about_ca_system_score_gemma":0.0018414389,"threshold_uncertainty_score":0.05312991},"labels":[],"label_agreement":null},{"id":"W4392619616","doi":"10.1145/3651159","title":"DeepMedFeature: An Accurate Feature Extraction and Drug-Drug Interaction Model for Clinical Text in Medical Informatics","year":2024,"lang":"en","type":"article","venue":"ACM Transactions on Asian and Low-Resource Language Information Processing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Brandon University","funders":"","keywords":"Computer science; Artificial intelligence; Informatics; Drug; Convolution (computer science); Feature extraction; F1 score; Feature vector; Mechanism (biology); Machine learning; Natural language processing; Artificial neural network; Medicine; Pharmacology","score_opus":0.016352322933853608,"score_gpt":0.3468787532546184,"score_spread":0.3305264303207648,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392619616","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07792418,0.004488471,0.87929296,0.0022409172,0.0004914761,0.0003781079,0.013634348,0.018798074,0.0027514088],"genre_scores_gemma":[0.55864143,0.0027557039,0.4016631,0.0011182677,0.00029997213,0.0006881921,0.02363595,0.00038625367,0.010811146],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9997441,0.000048264574,0.000034420056,0.0000758722,0.000061012084,0.00003629728],"domain_scores_gemma":[0.99969757,0.00015195001,0.000032853877,0.00003402032,0.000066034765,0.000017612654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00059735926,0.00084987504,0.00068745663,0.0010766268,0.0002152934,0.0005590567,0.00097331754,0.0010441542,0.0029910353],"category_scores_gemma":[0.0016358545,0.00026827815,0.0009739176,0.0009982787,0.00019785302,0.0013794143,0.0007677681,0.0013704541,0.0018267429],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006162754,0.000411297,0.0051649148,0.0004547529,0.00018158667,0.00051496015,0.00014243864,0.07662385,0.018832006,0.003505772,0.03228579,0.8612664],"study_design_scores_gemma":[0.000035322963,0.00015869865,0.0016935847,0.00003294793,0.000048611382,0.0002433689,0.000024563078,0.9772406,0.0061400756,0.0055389153,0.008823728,0.000019622352],"about_ca_topic_score_codex":0.005755677,"about_ca_topic_score_gemma":0.008267957,"teacher_disagreement_score":0.005755677,"about_ca_system_score_codex":0.00071831315,"about_ca_system_score_gemma":0.0011156491,"threshold_uncertainty_score":0.01144433},"labels":[],"label_agreement":null},{"id":"W4392656250","doi":"10.1186/s12874-024-02192-8","title":"Text analysis framework for identifying mutations among non-small cell lung cancer patients from laboratory data","year":2024,"lang":"en","type":"article","venue":"BMC Medical Research Methodology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Calgary Laboratory Services; University of Calgary","funders":"Canadian Institutes of Health Research","keywords":"Computer science; Lexical analysis; Natural language processing; Artificial intelligence; Information extraction; Syntax; Context (archaeology); Population; Information retrieval; Medicine","score_opus":0.35480033034478203,"score_gpt":0.5470634171707748,"score_spread":0.19226308682599275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392656250","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027047008,0.0017895901,0.9085647,0.0026889923,0.0002198853,0.0024930425,0.03506503,0.018085422,0.004046368],"genre_scores_gemma":[0.11969534,0.0007308894,0.81957793,0.000576472,0.00017578606,0.001216059,0.05560492,0.00025064423,0.0021719385],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99759537,0.0004057424,0.0004965479,0.00074980396,0.0006220591,0.00013050972],"domain_scores_gemma":[0.9949267,0.002838059,0.0005623327,0.00023398943,0.0012013192,0.00023755187],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002996126,0.0014365729,0.0009026892,0.010292404,0.0010285925,0.0024039962,0.0019101204,0.0010757609,0.0030067721],"category_scores_gemma":[0.009619177,0.00029498243,0.002209325,0.0039527086,0.0005843586,0.002430291,0.0013971686,0.0011535612,0.0016926884],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007287283,0.00070580264,0.028977968,0.0033769386,0.0005045165,0.0053781327,0.0021458096,0.038053043,0.026949286,0.02967497,0.059405573,0.80409926],"study_design_scores_gemma":[0.00016910005,0.00035934066,0.01799064,0.0007717803,0.0006555614,0.002618612,0.0016869099,0.7923891,0.023264872,0.073348746,0.08656633,0.00017901245],"about_ca_topic_score_codex":0.01103615,"about_ca_topic_score_gemma":0.015811868,"teacher_disagreement_score":0.01103615,"about_ca_system_score_codex":0.0013935169,"about_ca_system_score_gemma":0.0029674175,"threshold_uncertainty_score":0.021943867},"labels":[],"label_agreement":null},{"id":"W4392698591","doi":"10.1101/2024.03.10.24303884","title":"Applying Fast Healthcare Interoperability Resources (FHIR) for Pathogen Genomics at the Point of Care","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Providence Health Care; Simon Fraser University","funders":"Canadian Institutes of Health Research; Genome British Columbia; Genome Canada","keywords":"Interoperability; Computer science; Standardization; Data science; Terminology; Health care; Genomics; Knowledge management; World Wide Web; Biology","score_opus":0.023599818154072527,"score_gpt":0.29440868441822415,"score_spread":0.2708088662641516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392698591","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023447236,0.0006466335,0.86350214,0.021698575,0.000496922,0.0031362034,0.012524909,0.030130826,0.044416532],"genre_scores_gemma":[0.12519124,0.00036200657,0.8441992,0.0047336062,0.00014724256,0.0014316927,0.018117858,0.0024641436,0.0033529964],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9225307,0.046670265,0.0110945115,0.0039731045,0.012914126,0.0028171733],"domain_scores_gemma":[0.8095512,0.08651522,0.0110389115,0.057935365,0.031530287,0.003429079],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.121489994,0.0009779927,0.00087747374,0.006694703,0.0020631833,0.008310358,0.0043272492,0.0036634528,0.009339993],"category_scores_gemma":[0.13634036,0.0008418618,0.002033873,0.0041980557,0.0025094084,0.012998707,0.013308545,0.003549268,0.004897712],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009881238,0.00090577675,0.032283943,0.002975036,0.0003429885,0.0021296141,0.01119644,0.02750488,0.011486501,0.35535708,0.13736954,0.41746002],"study_design_scores_gemma":[0.00029688064,0.0005031631,0.012200286,0.0051942905,0.00020238745,0.0011352544,0.007037486,0.10651225,0.033668935,0.20696212,0.62571937,0.00056758634],"about_ca_topic_score_codex":0.017715858,"about_ca_topic_score_gemma":0.011903693,"teacher_disagreement_score":0.121489994,"about_ca_system_score_codex":0.0070783086,"about_ca_system_score_gemma":0.018476631,"threshold_uncertainty_score":0.64250815},"labels":[],"label_agreement":null},{"id":"W4392846105","doi":"10.1145/3625007.3629127","title":"KoExPubMed: A Tool for Effective and Customized Knowledge Extraction from PubMed","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Computer science; Extraction (chemistry); Information retrieval","score_opus":0.01754573588304488,"score_gpt":0.29965031217788207,"score_spread":0.2821045762948372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4392846105","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008696044,0.0026493717,0.4072475,0.0016486422,0.00034412942,0.0022594768,0.19662023,0.36816683,0.012367804],"genre_scores_gemma":[0.024900703,0.0026639488,0.8099185,0.00094853685,0.000194051,0.0032028907,0.13919124,0.013010356,0.005969771],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99817264,0.00041397408,0.0005546648,0.00032422625,0.00044269327,0.00009193841],"domain_scores_gemma":[0.99009645,0.007186386,0.0008671735,0.0007967473,0.0007237032,0.0003295342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003291321,0.002828636,0.001715792,0.01693993,0.0010462551,0.0034979857,0.0020522098,0.0012713192,0.03171483],"category_scores_gemma":[0.01584122,0.0012225145,0.0018419245,0.010906812,0.0004557256,0.0041972324,0.005462144,0.0013268056,0.015782947],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00092323794,0.00030108038,0.0049594934,0.020638583,0.0009447376,0.0038665244,0.0022145046,0.0027058388,0.021685923,0.0141664175,0.33945724,0.58813643],"study_design_scores_gemma":[0.00073583616,0.000285017,0.011208499,0.002998307,0.0004547215,0.004816492,0.0012237352,0.030093268,0.03134251,0.027832946,0.88847256,0.00053611264],"about_ca_topic_score_codex":0.0016056783,"about_ca_topic_score_gemma":0.0035574837,"teacher_disagreement_score":0.03171483,"about_ca_system_score_codex":0.0007247097,"about_ca_system_score_gemma":0.0032265547,"threshold_uncertainty_score":0.106096685},"labels":[],"label_agreement":null},{"id":"W4393032646","doi":"10.3819/ccbr.2024.190009","title":"The Future Is Computational Comparative Cognition","year":2024,"lang":"en","type":"article","venue":"Comparative Cognition & Behavior Reviews","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Comparative cognition; Animal cognition; Cognition; Cognitive science; Psychology; Animal behavior; Comparative psychology; Cognitive psychology; Neuroscience; Biology; Zoology","score_opus":0.10940314412105506,"score_gpt":0.4043020969445772,"score_spread":0.29489895282352213,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393032646","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009326996,0.19889963,0.19712327,0.4960256,0.0053075072,0.00006153154,0.00065789075,0.00083485903,0.09176278],"genre_scores_gemma":[0.43640608,0.16138595,0.2890725,0.06631946,0.018322721,0.00050276367,0.002035592,0.0010348621,0.024920082],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9946026,0.003174558,0.0002027948,0.00088417064,0.0009272531,0.00020861773],"domain_scores_gemma":[0.9747942,0.017417528,0.0006433542,0.0038555155,0.0024024327,0.00088706624],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.013185928,0.0006446344,0.0011814833,0.0024423827,0.0016156759,0.006769924,0.002423255,0.0031449415,0.011775456],"category_scores_gemma":[0.021008868,0.00034997737,0.0011507948,0.0021820466,0.013807087,0.024331966,0.0031992549,0.0050950227,0.0020669724],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027833677,0.000020529886,0.0006213899,0.00039916494,0.00005997867,0.000024396077,0.00054257025,0.0006805625,0.00016457612,0.9074485,0.022682358,0.067328036],"study_design_scores_gemma":[0.0000055954806,0.000008566725,0.0003247905,0.00014543161,0.000009181775,0.000033103806,0.00024047533,0.00080963573,0.000051614665,0.8516609,0.14669798,0.000012824769],"about_ca_topic_score_codex":0.0023732304,"about_ca_topic_score_gemma":0.0017804498,"teacher_disagreement_score":0.9868141,"about_ca_system_score_codex":0.0036363823,"about_ca_system_score_gemma":0.0038414798,"threshold_uncertainty_score":0.06973469},"labels":[],"label_agreement":null},{"id":"W4393306739","doi":"10.1093/genetics/iyae049","title":"Updates to the Alliance of Genome Resources central infrastructure","year":2024,"lang":"en","type":"article","venue":"Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":65,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Ontario Institute for Cancer Research","funders":"U.S. National Library of Medicine; University of Southern California; Wellcome Trust; National Human Genome Research Institute; Medical Research Council; European Molecular Biology Laboratory; NYU Grossman School of Medicine; University of Oregon; Harvard University; Eunice Kennedy Shriver National Institute of Child Health and Human Development; U.S. Department of Energy; California Institute of Technology; Lawrence Berkeley National Laboratory; National Heart, Lung, and Blood Institute","keywords":"Alliance; Biology; Genome; Genetics; Computational biology; Evolutionary biology; Gene; Political science","score_opus":0.0076573609430985005,"score_gpt":0.24890200716805125,"score_spread":0.24124464622495276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393306739","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0125414245,0.006071314,0.20789152,0.08373497,0.023294425,0.0036270698,0.14392152,0.25409502,0.26482275],"genre_scores_gemma":[0.018539025,0.0032950302,0.3360768,0.011911645,0.0036033974,0.0029921327,0.47921112,0.030375322,0.11399552],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9834487,0.0038627493,0.0019362239,0.0018974426,0.007386201,0.0014685831],"domain_scores_gemma":[0.9268617,0.006817304,0.0033385046,0.019680556,0.031246658,0.012055346],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.038156804,0.0013589012,0.0012531016,0.0091976905,0.0030138001,0.010586204,0.008099137,0.004413852,0.05929581],"category_scores_gemma":[0.0697767,0.0014779309,0.0011713716,0.01279633,0.0010431042,0.013352711,0.008902514,0.0070902756,0.06285024],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024376028,0.0001894569,0.0013078124,0.0003215059,0.000028119006,0.00013667677,0.00036738862,0.00056027365,0.002267301,0.015946584,0.8753673,0.10326371],"study_design_scores_gemma":[0.00003515138,0.000022318212,0.0005580578,0.00007011158,0.00001019059,0.000044538523,0.000046289817,0.00036690145,0.00060593017,0.0017035042,0.99650466,0.000032301446],"about_ca_topic_score_codex":0.018566003,"about_ca_topic_score_gemma":0.015295405,"teacher_disagreement_score":0.05929581,"about_ca_system_score_codex":0.005648145,"about_ca_system_score_gemma":0.016979428,"threshold_uncertainty_score":0.20179492},"labels":[],"label_agreement":null},{"id":"W4393325823","doi":"10.1007/978-3-031-47104-9","title":"Bayesian Filter Design for Computational Medicine","year":2024,"lang":"en","type":"book","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University; New York University; National Science Foundation","keywords":"Bayesian probability; Computer science; Artificial intelligence","score_opus":0.03896971679585074,"score_gpt":0.30469977260568964,"score_spread":0.2657300558098389,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393325823","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00024239748,0.0057180664,0.98223865,0.0010623235,0.00051676657,0.000020687718,0.00017161398,0.0004806176,0.009548752],"genre_scores_gemma":[0.040343326,0.023435447,0.85135096,0.0017467184,0.0021460298,0.0005167026,0.0015488844,0.0007939205,0.078117974],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99928635,0.0003010443,0.000033990502,0.00009352111,0.00026340692,0.000021790298],"domain_scores_gemma":[0.9985625,0.0010591466,0.00003721014,0.0001001536,0.00020914892,0.000031848045],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015337587,0.00097208726,0.00089397357,0.0005935268,0.000311179,0.0016227654,0.001039618,0.0015352214,0.018598262],"category_scores_gemma":[0.0057462268,0.00065577537,0.0007715523,0.0011584918,0.0008076964,0.001311669,0.0010516443,0.0025476364,0.007934792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000058127218,0.00004174212,0.00025698947,0.000527953,0.00013549117,0.000093953095,0.0001024045,0.10579021,0.0016978647,0.3185851,0.15202896,0.42068112],"study_design_scores_gemma":[0.0000230816,0.00004812238,0.00028663766,0.00021083668,0.00003528137,0.00014656768,0.000029662158,0.37045488,0.0009655849,0.45281535,0.17494605,0.000038010476],"about_ca_topic_score_codex":0.00221224,"about_ca_topic_score_gemma":0.0024645596,"teacher_disagreement_score":0.018598262,"about_ca_system_score_codex":0.0008070551,"about_ca_system_score_gemma":0.0008041851,"threshold_uncertainty_score":0.062217355},"labels":[],"label_agreement":null},{"id":"W4393390454","doi":"10.1002/ca.24162","title":"<scp>TA2Viewer</scp>: A web‐based browser for <i>Terminologia Anatomica</i> and online anatomical knowledge","year":2024,"lang":"en","type":"article","venue":"Clinical Anatomy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"National Institute of Biomedical Imaging and Bioengineering; National Cancer Institute; National Institutes of Health; Leidos","keywords":"Terminology; Computer science; World Wide Web; Listing (finance); Information retrieval; Unified Medical Language System; Medical terminology; Hierarchy; Medicine; Linguistics","score_opus":0.04467729164823545,"score_gpt":0.38960579548204805,"score_spread":0.34492850383381257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393390454","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018696395,0.0015201003,0.098784186,0.0019943095,0.0019979684,0.0009553529,0.33024713,0.3369385,0.22569285],"genre_scores_gemma":[0.016675185,0.002015389,0.13118601,0.0049221297,0.0015763025,0.0021040738,0.49308693,0.15969935,0.18873468],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992749,0.0001148364,0.00009977635,0.00011080832,0.00030969095,0.00008991453],"domain_scores_gemma":[0.99438983,0.002269432,0.00033932785,0.00074399146,0.0015363568,0.0007210077],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0014005404,0.0020630744,0.0011707463,0.0058210413,0.0010524797,0.0048826453,0.002617377,0.0021822366,0.4631719],"category_scores_gemma":[0.007996403,0.0008648007,0.0011488725,0.006158535,0.0008192397,0.007536234,0.0056038396,0.0029752136,0.30062088],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048690308,0.000013589104,0.00013081901,0.0005337768,0.0000128781685,0.0001753601,0.00020952518,0.0000679603,0.0011493418,0.0024851344,0.97366107,0.021511912],"study_design_scores_gemma":[0.000038014612,0.0000102193435,0.00077366276,0.00023003054,0.000011548469,0.0004140147,0.00011163488,0.0008692479,0.0012656049,0.004406401,0.9918243,0.000045427467],"about_ca_topic_score_codex":0.0060922066,"about_ca_topic_score_gemma":0.013093681,"teacher_disagreement_score":0.4631719,"about_ca_system_score_codex":0.0011639578,"about_ca_system_score_gemma":0.0022462606,"threshold_uncertainty_score":0.7657201},"labels":[],"label_agreement":null},{"id":"W4393490244","doi":"10.5281/zenodo.4007064","title":"Results of Neuron Phenotype Ontology competency queries","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Addiction and Mental Health","funders":"","keywords":"Phenotype; Ontology; Computer science; Computational biology; Information retrieval; Biology; Gene; Genetics; Philosophy; Epistemology","score_opus":0.031615937000210606,"score_gpt":0.26159619732734307,"score_spread":0.22998026032713248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393490244","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008912468,0.00049269944,0.000980942,0.00030158446,0.00006233448,0.000045204175,0.98233944,0.0018362518,0.0050290753],"genre_scores_gemma":[0.006780371,0.0001337217,0.002216207,0.00010330522,0.000006526432,0.000067050205,0.98874986,0.00014018056,0.0018029031],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975707,0.00034099034,0.00023687202,0.0007078724,0.00089118653,0.0002524059],"domain_scores_gemma":[0.99736696,0.0011711263,0.00013131849,0.000540229,0.00062639767,0.00016408222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001679658,0.0017743729,0.00082630524,0.0033778097,0.0008231833,0.0020392209,0.0011440265,0.0014865071,0.017137486],"category_scores_gemma":[0.007950892,0.00026203037,0.0014653074,0.002793549,0.00046006395,0.0011044593,0.0015155568,0.001172255,0.017025642],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044340687,0.00021742225,0.010400778,0.00137304,0.00013123217,0.00026659158,0.000119709075,0.0036319161,0.0027212861,0.0029063558,0.9495289,0.02825929],"study_design_scores_gemma":[0.00041283673,0.000111156216,0.042607933,0.00043988132,0.0001596984,0.00086665654,0.00080134545,0.015046214,0.011598037,0.009024849,0.9188323,0.000099112855],"about_ca_topic_score_codex":0.022680486,"about_ca_topic_score_gemma":0.038341157,"teacher_disagreement_score":0.022680486,"about_ca_system_score_codex":0.0022553222,"about_ca_system_score_gemma":0.0021677127,"threshold_uncertainty_score":0.057330668},"labels":[],"label_agreement":null},{"id":"W4393537727","doi":"10.5281/zenodo.2039605","title":"All simulated data for clonealign paper","year":2018,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science","score_opus":0.07477758024516973,"score_gpt":0.3178399315028199,"score_spread":0.2430623512576502,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393537727","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028910011,0.00044448584,0.0055418992,0.0008640109,0.00054388697,0.00032002386,0.953342,0.005088028,0.004945627],"genre_scores_gemma":[0.021647137,0.000114106224,0.0064863293,0.00027280248,0.000029138808,0.0005220374,0.96667796,0.00037604079,0.0038743704],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99795157,0.00050813094,0.00014395299,0.0007008396,0.000519011,0.0001765048],"domain_scores_gemma":[0.9945891,0.0019592321,0.00019906058,0.0021174,0.00085460837,0.00028062324],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002858405,0.0016076702,0.0008220585,0.0014858298,0.0007864584,0.001576573,0.0021345539,0.0018533936,0.022177989],"category_scores_gemma":[0.010791905,0.0005587025,0.0015503383,0.0020257786,0.00081740343,0.0011150042,0.0011843477,0.0023868333,0.020121915],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010851208,0.00066872383,0.015113,0.00070274604,0.0002716051,0.00014753104,0.00009331939,0.03351209,0.0016312415,0.0032568562,0.9205979,0.022919862],"study_design_scores_gemma":[0.0017232891,0.00038983082,0.022957265,0.00029290197,0.00019302708,0.0007297867,0.00034694446,0.090277135,0.014142044,0.01228436,0.8565233,0.00014011396],"about_ca_topic_score_codex":0.008027433,"about_ca_topic_score_gemma":0.01722736,"teacher_disagreement_score":0.022177989,"about_ca_system_score_codex":0.0017404865,"about_ca_system_score_gemma":0.0018891133,"threshold_uncertainty_score":0.07419276},"labels":[],"label_agreement":null},{"id":"W4393547873","doi":"10.5281/zenodo.1227313","title":"Data Associated With \"A Collaborative Filtering Based Approach To Biomedical Knowledge Discovery\"","year":2018,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science; Data science; Knowledge extraction; Collaborative filtering; Information retrieval; Data mining; Recommender system","score_opus":0.05704799744159094,"score_gpt":0.2958549076893579,"score_spread":0.23880691024776696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393547873","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018834737,0.00015885953,0.000446539,0.0001983787,0.00006580022,0.00006656404,0.99579537,0.0006626294,0.0007223113],"genre_scores_gemma":[0.0017787921,0.000069742484,0.0011524962,0.00005339158,0.000010455573,0.00012576283,0.9962352,0.000036790843,0.00053740054],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99835783,0.00026534154,0.00026677133,0.00042718532,0.00049980666,0.0001829564],"domain_scores_gemma":[0.996421,0.0011330892,0.00042001723,0.0007985728,0.0008625248,0.00036491032],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015541273,0.0019458791,0.0010999658,0.0035284243,0.000718711,0.001576053,0.0025852039,0.002269893,0.02065389],"category_scores_gemma":[0.007703057,0.00047332453,0.0015327133,0.005178499,0.00056091655,0.0010817683,0.0016265924,0.001758591,0.01968685],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000564158,0.00025156158,0.0050923736,0.002329196,0.00013232893,0.0002463743,0.00010394307,0.0017292955,0.0013037625,0.0010715595,0.9741279,0.013047552],"study_design_scores_gemma":[0.00061606703,0.00015113348,0.01831618,0.00039789407,0.000109428904,0.00056940113,0.00024590295,0.0028336004,0.0041034142,0.0016938873,0.9708911,0.00007190293],"about_ca_topic_score_codex":0.0152004585,"about_ca_topic_score_gemma":0.025757156,"teacher_disagreement_score":0.02065389,"about_ca_system_score_codex":0.0015847379,"about_ca_system_score_gemma":0.0026459605,"threshold_uncertainty_score":0.06909418},"labels":[],"label_agreement":null},{"id":"W4393557647","doi":"10.5281/zenodo.4007065","title":"Results of Neuron Phenotype Ontology competency queries","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Addiction and Mental Health","funders":"","keywords":"Phenotype; Ontology; Computer science; Computational biology; Information retrieval; Artificial intelligence; Biology; Genetics; Gene; Philosophy; Epistemology","score_opus":0.031615937000210606,"score_gpt":0.26159619732734307,"score_spread":0.22998026032713248,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393557647","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008912468,0.00049269944,0.000980942,0.00030158446,0.00006233448,0.000045204175,0.98233944,0.0018362518,0.0050290753],"genre_scores_gemma":[0.006780371,0.0001337217,0.002216207,0.00010330522,0.000006526432,0.000067050205,0.98874986,0.00014018056,0.0018029031],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975707,0.00034099034,0.00023687202,0.0007078724,0.00089118653,0.0002524059],"domain_scores_gemma":[0.99736696,0.0011711263,0.00013131849,0.000540229,0.00062639767,0.00016408222],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001679658,0.0017743729,0.00082630524,0.0033778097,0.0008231833,0.0020392209,0.0011440265,0.0014865071,0.017137486],"category_scores_gemma":[0.007950892,0.00026203037,0.0014653074,0.002793549,0.00046006395,0.0011044593,0.0015155568,0.001172255,0.017025642],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044340687,0.00021742225,0.010400778,0.00137304,0.00013123217,0.00026659158,0.000119709075,0.0036319161,0.0027212861,0.0029063558,0.9495289,0.02825929],"study_design_scores_gemma":[0.00041283673,0.000111156216,0.042607933,0.00043988132,0.0001596984,0.00086665654,0.00080134545,0.015046214,0.011598037,0.009024849,0.9188323,0.000099112855],"about_ca_topic_score_codex":0.022680486,"about_ca_topic_score_gemma":0.038341157,"teacher_disagreement_score":0.022680486,"about_ca_system_score_codex":0.0022553222,"about_ca_system_score_gemma":0.0021677127,"threshold_uncertainty_score":0.057330668},"labels":[],"label_agreement":null},{"id":"W4393612922","doi":"10.5281/zenodo.8170024","title":"OpenAlex Author Name Disambiguation V3 Initial Clusters","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"OpenAlex","funders":"","keywords":"Natural language processing; Computer science; Information retrieval; Linguistics; Artificial intelligence; Philosophy","score_opus":0.06265774508346136,"score_gpt":0.3217900060954892,"score_spread":0.2591322610120278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393612922","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001812901,0.00026410943,0.0011616667,0.00016011106,0.00012206527,0.00013023295,0.9869408,0.005344713,0.004063274],"genre_scores_gemma":[0.0009938566,0.000058135647,0.002697687,0.000051683248,0.000014038511,0.00018397371,0.99430746,0.00034207437,0.0013511382],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99738485,0.00030930364,0.0002661885,0.0010526724,0.0005376166,0.00044928837],"domain_scores_gemma":[0.99734086,0.0005803564,0.00017222349,0.00082487083,0.00081226445,0.00026938986],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.001857242,0.0036401334,0.0018063035,0.007164037,0.0026643353,0.004455758,0.0035099054,0.0026665118,0.047998562],"category_scores_gemma":[0.0065701385,0.00089967996,0.0024056414,0.007104338,0.00091747957,0.0022258128,0.0037584908,0.0026189014,0.099639185],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001759253,0.000059920116,0.0012957226,0.0005675721,0.000038826307,0.000073527284,0.00008880008,0.0005129344,0.00052249525,0.001277793,0.9878292,0.0075572524],"study_design_scores_gemma":[0.0002205354,0.000026462847,0.0029803813,0.00022553842,0.000053353895,0.00015641001,0.00027200487,0.0014791193,0.002212759,0.0028087765,0.9895139,0.000050758128],"about_ca_topic_score_codex":0.015069076,"about_ca_topic_score_gemma":0.035150304,"teacher_disagreement_score":0.9981428,"about_ca_system_score_codex":0.0023123082,"about_ca_system_score_gemma":0.0043976526,"threshold_uncertainty_score":0.16057122},"labels":[],"label_agreement":null},{"id":"W4393657389","doi":"10.5281/zenodo.8170023","title":"OpenAlex Author Name Disambiguation V3 Initial Clusters","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"OpenAlex","funders":"","keywords":"Natural language processing; Computer science; Artificial intelligence","score_opus":0.06265774508346136,"score_gpt":0.3217900060954892,"score_spread":0.2591322610120278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393657389","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001812901,0.00026410943,0.0011616667,0.00016011106,0.00012206527,0.00013023295,0.9869408,0.005344713,0.004063274],"genre_scores_gemma":[0.0009938566,0.000058135647,0.002697687,0.000051683248,0.000014038511,0.00018397371,0.99430746,0.00034207437,0.0013511382],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99738485,0.00030930364,0.0002661885,0.0010526724,0.0005376166,0.00044928837],"domain_scores_gemma":[0.99734086,0.0005803564,0.00017222349,0.00082487083,0.00081226445,0.00026938986],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.001857242,0.0036401334,0.0018063035,0.007164037,0.0026643353,0.004455758,0.0035099054,0.0026665118,0.047998562],"category_scores_gemma":[0.0065701385,0.00089967996,0.0024056414,0.007104338,0.00091747957,0.0022258128,0.0037584908,0.0026189014,0.099639185],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001759253,0.000059920116,0.0012957226,0.0005675721,0.000038826307,0.000073527284,0.00008880008,0.0005129344,0.00052249525,0.001277793,0.9878292,0.0075572524],"study_design_scores_gemma":[0.0002205354,0.000026462847,0.0029803813,0.00022553842,0.000053353895,0.00015641001,0.00027200487,0.0014791193,0.002212759,0.0028087765,0.9895139,0.000050758128],"about_ca_topic_score_codex":0.015069076,"about_ca_topic_score_gemma":0.035150304,"teacher_disagreement_score":0.9981428,"about_ca_system_score_codex":0.0023123082,"about_ca_system_score_gemma":0.0043976526,"threshold_uncertainty_score":0.16057122},"labels":[],"label_agreement":null},{"id":"W4393662871","doi":"10.5281/zenodo.7392043","title":"Supplementary material for a usability evaluation of a semantic search for biological datasets","year":2022,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Usability; Computer science; Information retrieval; World Wide Web; Human–computer interaction","score_opus":0.08542570881769637,"score_gpt":0.3367827754308616,"score_spread":0.25135706661316526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393662871","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010680817,0.00024795858,0.005169303,0.00032364673,0.00014602196,0.0006583154,0.96782213,0.007324758,0.0076269624],"genre_scores_gemma":[0.01036271,0.00008365725,0.011868953,0.00025206088,0.00001855871,0.0019069879,0.9710708,0.0010044248,0.0034318613],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99701285,0.0009770371,0.000477923,0.00049115036,0.00085966353,0.0001813396],"domain_scores_gemma":[0.9864305,0.00839482,0.00040331384,0.0016261695,0.0026801662,0.00046505887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037908508,0.0014955905,0.0006774273,0.0034618392,0.00087822817,0.0016590648,0.0009906476,0.00089614576,0.099816106],"category_scores_gemma":[0.015584064,0.0003429351,0.0008125555,0.003227245,0.00033087368,0.001149456,0.0016124725,0.0008384471,0.04229102],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038183073,0.0005886687,0.004778359,0.0018002613,0.000084616804,0.00013747178,0.000235777,0.0011148867,0.0011664374,0.0010661503,0.96218956,0.026455883],"study_design_scores_gemma":[0.000826477,0.00038660524,0.03691099,0.0006062456,0.00012651696,0.0006442438,0.0008828805,0.007826259,0.0064308764,0.0049800775,0.9402536,0.00012522534],"about_ca_topic_score_codex":0.007851419,"about_ca_topic_score_gemma":0.02505152,"teacher_disagreement_score":0.099816106,"about_ca_system_score_codex":0.0014737493,"about_ca_system_score_gemma":0.0012543071,"threshold_uncertainty_score":0.33391815},"labels":[],"label_agreement":null},{"id":"W4393663063","doi":"10.5281/zenodo.7539374","title":"The Nova Scotia Disease Knowledge Graph","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Cape Breton University","funders":"","keywords":"Nova scotia; Nova (rocket); Graph; Computer science; Geography; Combinatorics; Mathematics; Engineering; Archaeology; Aeronautics","score_opus":0.0394137044001122,"score_gpt":0.2897344116835551,"score_spread":0.2503207072834429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393663063","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059809345,0.00045429158,0.00044415853,0.00042563243,0.000039174764,0.000054632485,0.98662853,0.00041782012,0.00555489],"genre_scores_gemma":[0.009003881,0.00027482514,0.0012212215,0.00011681594,0.0000074569048,0.000056979778,0.9870401,0.000036319136,0.0022424886],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995721,0.000055283475,0.000034443947,0.00014405204,0.000116870484,0.00007725845],"domain_scores_gemma":[0.99898094,0.00026199105,0.0000857754,0.00021158146,0.00027124753,0.00018843944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041809323,0.00083525426,0.00048567288,0.0042973207,0.0010134391,0.0010916522,0.0012391,0.0009369624,0.014133352],"category_scores_gemma":[0.0025692375,0.00032282592,0.00059722323,0.005736985,0.0004446365,0.00046049547,0.0012088692,0.0009677588,0.006852093],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023182349,0.00007682703,0.0103327725,0.0010288601,0.000098444114,0.00074036646,0.00019080777,0.0027925717,0.0011262314,0.0040901625,0.9630391,0.016252065],"study_design_scores_gemma":[0.00013802444,0.000026526783,0.03786877,0.00041980357,0.000082325605,0.00060966995,0.00029535237,0.0028714945,0.0011473469,0.0029499582,0.953557,0.00003372086],"about_ca_topic_score_codex":0.4589241,"about_ca_topic_score_gemma":0.6103371,"teacher_disagreement_score":0.5410759,"about_ca_system_score_codex":0.004112629,"about_ca_system_score_gemma":0.0050680274,"threshold_uncertainty_score":0.91250575},"labels":[],"label_agreement":null},{"id":"W4393701562","doi":"10.5281/zenodo.5517235","title":"The Nova Scotia Disease Knowledge Graph","year":2023,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Cape Breton University","funders":"","keywords":"Nova scotia; Nova (rocket); Graph; Geography; Computer science; Engineering; Archaeology; Theoretical computer science; Aeronautics","score_opus":0.0394137044001122,"score_gpt":0.2897344116835551,"score_spread":0.2503207072834429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393701562","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0059809345,0.00045429158,0.00044415853,0.00042563243,0.000039174764,0.000054632485,0.98662853,0.00041782012,0.00555489],"genre_scores_gemma":[0.009003881,0.00027482514,0.0012212215,0.00011681594,0.0000074569048,0.000056979778,0.9870401,0.000036319136,0.0022424886],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995721,0.000055283475,0.000034443947,0.00014405204,0.000116870484,0.00007725845],"domain_scores_gemma":[0.99898094,0.00026199105,0.0000857754,0.00021158146,0.00027124753,0.00018843944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00041809323,0.00083525426,0.00048567288,0.0042973207,0.0010134391,0.0010916522,0.0012391,0.0009369624,0.014133352],"category_scores_gemma":[0.0025692375,0.00032282592,0.00059722323,0.005736985,0.0004446365,0.00046049547,0.0012088692,0.0009677588,0.006852093],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023182349,0.00007682703,0.0103327725,0.0010288601,0.000098444114,0.00074036646,0.00019080777,0.0027925717,0.0011262314,0.0040901625,0.9630391,0.016252065],"study_design_scores_gemma":[0.00013802444,0.000026526783,0.03786877,0.00041980357,0.000082325605,0.00060966995,0.00029535237,0.0028714945,0.0011473469,0.0029499582,0.953557,0.00003372086],"about_ca_topic_score_codex":0.4589241,"about_ca_topic_score_gemma":0.6103371,"teacher_disagreement_score":0.5410759,"about_ca_system_score_codex":0.004112629,"about_ca_system_score_gemma":0.0050680274,"threshold_uncertainty_score":0.91250575},"labels":[],"label_agreement":null},{"id":"W4393841825","doi":"10.5281/zenodo.7388037","title":"Supplementary material for a usability evaluation of a semantic search for biological datasets","year":2022,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Usability; Computer science; Information retrieval; World Wide Web; Human–computer interaction","score_opus":0.08542570881769637,"score_gpt":0.3367827754308616,"score_spread":0.25135706661316526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393841825","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010680817,0.00024795858,0.005169303,0.00032364673,0.00014602196,0.0006583154,0.96782213,0.007324758,0.0076269624],"genre_scores_gemma":[0.01036271,0.00008365725,0.011868953,0.00025206088,0.00001855871,0.0019069879,0.9710708,0.0010044248,0.0034318613],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99701285,0.0009770371,0.000477923,0.00049115036,0.00085966353,0.0001813396],"domain_scores_gemma":[0.9864305,0.00839482,0.00040331384,0.0016261695,0.0026801662,0.00046505887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0037908508,0.0014955905,0.0006774273,0.0034618392,0.00087822817,0.0016590648,0.0009906476,0.00089614576,0.099816106],"category_scores_gemma":[0.015584064,0.0003429351,0.0008125555,0.003227245,0.00033087368,0.001149456,0.0016124725,0.0008384471,0.04229102],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038183073,0.0005886687,0.004778359,0.0018002613,0.000084616804,0.00013747178,0.000235777,0.0011148867,0.0011664374,0.0010661503,0.96218956,0.026455883],"study_design_scores_gemma":[0.000826477,0.00038660524,0.03691099,0.0006062456,0.00012651696,0.0006442438,0.0008828805,0.007826259,0.0064308764,0.0049800775,0.9402536,0.00012522534],"about_ca_topic_score_codex":0.007851419,"about_ca_topic_score_gemma":0.02505152,"teacher_disagreement_score":0.099816106,"about_ca_system_score_codex":0.0014737493,"about_ca_system_score_gemma":0.0012543071,"threshold_uncertainty_score":0.33391815},"labels":[],"label_agreement":null},{"id":"W4393891061","doi":"10.5281/zenodo.3405609","title":"BioWordlists","year":2019,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Computer science","score_opus":0.04771689702469021,"score_gpt":0.3102005375547082,"score_spread":0.262483640530018,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4393891061","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00051171385,0.0005236136,0.008039625,0.0009637425,0.00076144503,0.00028014858,0.86813974,0.09037964,0.03040029],"genre_scores_gemma":[0.0016419641,0.00054576655,0.014761867,0.00092045125,0.00019965605,0.00076372735,0.90776783,0.04457923,0.028819492],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9973972,0.00043983513,0.0005114107,0.000578095,0.00081752345,0.00025599852],"domain_scores_gemma":[0.9903998,0.0038376937,0.0009991775,0.0012197944,0.0029227103,0.0006208799],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0033190837,0.0039774915,0.0021377383,0.010600713,0.0017347207,0.007340516,0.003254633,0.0027643798,0.5836663],"category_scores_gemma":[0.01884672,0.0021107632,0.0020477157,0.01153446,0.00063693983,0.0073653236,0.005030136,0.0023961572,0.60247356],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014985309,0.00001782759,0.00019702908,0.0011850178,0.00001829935,0.000060505718,0.00008478358,0.00009237156,0.0004679081,0.0015727947,0.9698617,0.026292017],"study_design_scores_gemma":[0.00007068128,0.000027575616,0.00038438884,0.00028057015,0.00001304824,0.000094151736,0.00006884328,0.00016847209,0.00085603836,0.0023753447,0.9956279,0.00003299548],"about_ca_topic_score_codex":0.002742273,"about_ca_topic_score_gemma":0.002304556,"teacher_disagreement_score":0.5836663,"about_ca_system_score_codex":0.0021626246,"about_ca_system_score_gemma":0.0034979708,"threshold_uncertainty_score":0.5938494},"labels":[],"label_agreement":null},{"id":"W4394053506","doi":"10.5281/zenodo.7388038","title":"Supplementary material for a usability evaluation of a semantic search for biological datasets","year":2022,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Usability; Computer science; Semantic search; Information retrieval; World Wide Web; Semantic Web; Human–computer interaction","score_opus":0.08542570881769637,"score_gpt":0.3367827754308616,"score_spread":0.25135706661316526,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394053506","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008163221,0.00020508433,0.0039087506,0.00027298916,0.00012595538,0.0005098059,0.9744634,0.0058124326,0.0065383576],"genre_scores_gemma":[0.0077765626,0.00006704216,0.008639811,0.00019982421,0.000015502266,0.0014889018,0.9781939,0.0007252125,0.0028932178],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9972301,0.000867952,0.00043204194,0.00048041483,0.0008088736,0.00018069283],"domain_scores_gemma":[0.9884524,0.0068291267,0.0003762932,0.0014988679,0.0024036004,0.0004398363],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0034401533,0.0015759186,0.0006889842,0.0035509516,0.00090761366,0.0016373432,0.0010362614,0.0009362543,0.10279013],"category_scores_gemma":[0.013997169,0.00035273912,0.00082738465,0.0032553647,0.00033902176,0.0011405043,0.0016123832,0.00086846866,0.046252612],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003139017,0.00049173826,0.0041325963,0.0014579337,0.000070776456,0.000118528485,0.00017683241,0.000998897,0.0009209258,0.0009454413,0.9692096,0.02116273],"study_design_scores_gemma":[0.00075501035,0.0003226678,0.032648876,0.0005155088,0.00010742353,0.00059677754,0.00074316113,0.0069208858,0.005353271,0.0046507786,0.94727576,0.00010999283],"about_ca_topic_score_codex":0.008638292,"about_ca_topic_score_gemma":0.02839031,"teacher_disagreement_score":0.99655986,"about_ca_system_score_codex":0.0015335607,"about_ca_system_score_gemma":0.0012577106,"threshold_uncertainty_score":0.34386724},"labels":[],"label_agreement":null},{"id":"W4394131026","doi":"10.6084/m9.figshare.19401737","title":"Additional file 3 of The Xenopus phenotype ontology: bridging model organism phenotype data to human health and development","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Phenotype; Bridging (networking); Xenopus; Organism; Model organism; Computational biology; Ontology; Computer science; Biology; Genetics; Gene; Computer security","score_opus":0.08255026807067962,"score_gpt":0.3075142445151045,"score_spread":0.2249639764444249,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394131026","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00005599423,0.000010930832,0.000056114295,0.000026230156,0.0000056526283,0.000010333876,0.9994779,0.00012107703,0.00023575709],"genre_scores_gemma":[0.00068448496,0.00003297917,0.0005212127,0.00006205736,0.000005725244,0.00015496144,0.99776804,0.000095627074,0.0006748438],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990922,0.00011479471,0.00015013918,0.0003014782,0.00020723771,0.00013414565],"domain_scores_gemma":[0.99447083,0.003369152,0.00048699082,0.0005902727,0.0007750372,0.00030764355],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0014641742,0.0018921889,0.0013455766,0.0035563293,0.0008609507,0.0019734032,0.002232609,0.0017653721,0.3765457],"category_scores_gemma":[0.008845724,0.0006866522,0.0013589017,0.0050013047,0.0004997775,0.0016934501,0.0017538547,0.0015914494,0.087306194],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016770585,0.000057910132,0.002238997,0.0022098552,0.00005944002,0.000069581554,0.000054214048,0.00042097364,0.00032634946,0.0008307043,0.9903106,0.0032537782],"study_design_scores_gemma":[0.0008736068,0.000053557662,0.0102801295,0.00087073306,0.000115930394,0.00026135508,0.00019474713,0.000783622,0.0011018917,0.0035354295,0.9818649,0.000064109176],"about_ca_topic_score_codex":0.014304695,"about_ca_topic_score_gemma":0.026637243,"teacher_disagreement_score":0.3765457,"about_ca_system_score_codex":0.0017890041,"about_ca_system_score_gemma":0.0022470132,"threshold_uncertainty_score":0.88928187},"labels":[],"label_agreement":null},{"id":"W4394329559","doi":"10.6084/m9.figshare.20175971","title":"Additional file 1 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"dataset","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Computer science; Information retrieval; Data science; World Wide Web","score_opus":0.08444570121034899,"score_gpt":0.39317620636855255,"score_spread":0.30873050515820355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394329559","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007325741,0.000009769107,0.0000565553,0.000030007863,0.0000060311672,0.000018287965,0.9994537,0.00006363251,0.00028875747],"genre_scores_gemma":[0.00061862596,0.000029059524,0.0005092108,0.00007323156,0.000007571635,0.00020529487,0.9977775,0.000057599078,0.000721883],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998896,0.00013422141,0.00021210246,0.0003446046,0.00023992425,0.00017316813],"domain_scores_gemma":[0.99413526,0.0033006205,0.000609272,0.0006081731,0.00095792627,0.00038868975],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016983398,0.0014096885,0.0012584591,0.0033553438,0.00088926277,0.0019707885,0.0019940254,0.0015773305,0.3425673],"category_scores_gemma":[0.00978098,0.0007242696,0.0010477583,0.006444935,0.00055289327,0.0017238247,0.001679586,0.0018058118,0.09756594],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014015927,0.000078941615,0.0019481927,0.0014788742,0.00003955961,0.000049985436,0.00004263043,0.00034727342,0.00020425695,0.0009751703,0.9915182,0.0031767557],"study_design_scores_gemma":[0.0012362668,0.00006764322,0.010552409,0.0007473375,0.00008949911,0.00026205566,0.00023009583,0.0006906708,0.0009554668,0.0050166473,0.9800913,0.000060718306],"about_ca_topic_score_codex":0.011954731,"about_ca_topic_score_gemma":0.021190073,"teacher_disagreement_score":0.3425673,"about_ca_system_score_codex":0.0022212532,"about_ca_system_score_gemma":0.0033774038,"threshold_uncertainty_score":0.93774796},"labels":[],"label_agreement":null},{"id":"W4394410433","doi":"10.6084/m9.figshare.19401731","title":"Additional file 1 of The Xenopus phenotype ontology: bridging model organism phenotype data to human health and development","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Phenotype; Bridging (networking); Xenopus; Model organism; Ontology; Computational biology; Organism; Biology; Computer science; Genetics; Bioinformatics; Gene; Computer security; Philosophy","score_opus":0.0839332524157449,"score_gpt":0.3079487255259383,"score_spread":0.22401547311019343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394410433","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000700531,0.000013197478,0.000074541975,0.000029932735,0.000005505212,0.0000130921535,0.9994246,0.00012427058,0.0002448262],"genre_scores_gemma":[0.00079668564,0.00004180736,0.00069304113,0.000075627664,0.000006158288,0.0001978668,0.9974222,0.00009864262,0.0006680123],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990976,0.0001219745,0.00014657585,0.00030643487,0.00020089665,0.00012647842],"domain_scores_gemma":[0.99319345,0.004386138,0.00052038365,0.0007043181,0.0008493017,0.00034651556],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0014961404,0.0018733114,0.0013549112,0.003543911,0.0008718836,0.0019286916,0.0023942853,0.0019209462,0.37716362],"category_scores_gemma":[0.010499698,0.0007325191,0.0013565161,0.005363828,0.00054769823,0.0017308153,0.0017172864,0.0016670949,0.08728433],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017125135,0.000076419725,0.002286871,0.002669382,0.000059240196,0.00007099501,0.000057038804,0.0005524854,0.0003126263,0.0010599033,0.9886828,0.004001048],"study_design_scores_gemma":[0.001148199,0.00007095162,0.010782363,0.0010104754,0.00012758144,0.00030459202,0.00019360844,0.0011115205,0.0011064041,0.0048588905,0.9792138,0.000071612514],"about_ca_topic_score_codex":0.012244147,"about_ca_topic_score_gemma":0.02318945,"teacher_disagreement_score":0.37716362,"about_ca_system_score_codex":0.0018841954,"about_ca_system_score_gemma":0.0026271383,"threshold_uncertainty_score":0.8884005},"labels":[],"label_agreement":null},{"id":"W4394544620","doi":"10.6084/m9.figshare.19401734","title":"Additional file 2 of The Xenopus phenotype ontology: bridging model organism phenotype data to human health and development","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Phenotype; Bridging (networking); Xenopus; Organism; Computational biology; Model organism; Biology; Ontology; Computer science; Genetics; Gene; Computer security; Philosophy","score_opus":0.0813036634809764,"score_gpt":0.3073642430765197,"score_spread":0.22606057959554332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394544620","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006132533,0.000010405146,0.000061529136,0.000028884006,0.0000057921698,0.000013315951,0.99945027,0.00010890153,0.00025953868],"genre_scores_gemma":[0.0007730682,0.000034247623,0.0006466525,0.00007923351,0.000007403185,0.00021325766,0.9973355,0.000108799584,0.0008018084],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990871,0.00012079998,0.00014590796,0.0003120541,0.00019555913,0.00013853126],"domain_scores_gemma":[0.9933588,0.0041317125,0.0005251921,0.00072040304,0.0008942261,0.00036968692],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015005729,0.0017643068,0.0014024142,0.0034528999,0.000933682,0.0020752372,0.0023440935,0.0017677053,0.44274947],"category_scores_gemma":[0.010142353,0.0007384986,0.0013346112,0.005282158,0.000539701,0.0017348643,0.0016771739,0.0016519506,0.09771498],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015903651,0.00006878257,0.0020813367,0.0022163354,0.000050580005,0.00006216044,0.00005057745,0.0003971422,0.00025985963,0.0008245192,0.9901472,0.0036826213],"study_design_scores_gemma":[0.0011213858,0.000066180924,0.010513504,0.0009172709,0.00011443961,0.00025987395,0.00019148146,0.0008180577,0.0009990195,0.0042983503,0.9806331,0.00006749124],"about_ca_topic_score_codex":0.013161523,"about_ca_topic_score_gemma":0.023970474,"teacher_disagreement_score":0.44274947,"about_ca_system_score_codex":0.0017403765,"about_ca_system_score_gemma":0.0024086873,"threshold_uncertainty_score":0.79485023},"labels":[],"label_agreement":null},{"id":"W4394601735","doi":"10.1212/wnl.0000000000205725","title":"Exploring Knowledgeable Informant Reporting on the Earliest Changes of Parkinson’s Disease (P4-3.010)","year":2024,"lang":"en","type":"article","venue":"Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Movement Disorders","funders":"","keywords":"Parkinson's disease; Disease; Psychology; Value (mathematics); Medicine; Gerontology; Clinical psychology; Internal medicine; Statistics; Mathematics","score_opus":0.11605255193030987,"score_gpt":0.28924236169966167,"score_spread":0.17318980976935178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394601735","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8429537,0.0020886364,0.06278024,0.0054903743,0.00021318805,0.0005161845,0.06970501,0.000895656,0.015356971],"genre_scores_gemma":[0.913476,0.0007152037,0.0530282,0.0004448505,0.00005298558,0.0002162324,0.030927522,0.00006519545,0.001073822],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99642134,0.0014640737,0.00045406778,0.0006030772,0.0008944388,0.0001629943],"domain_scores_gemma":[0.9391425,0.046641055,0.0062405113,0.0026431007,0.0046608476,0.0006719539],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075145415,0.0003451892,0.00032356416,0.0035087445,0.00040332368,0.0017618285,0.0007062464,0.0007190247,0.0023917947],"category_scores_gemma":[0.05523906,0.0001270646,0.000634347,0.0029344007,0.00033742745,0.0027192726,0.0015252605,0.00066582,0.0005450446],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025979057,0.0006609004,0.6691468,0.0027037878,0.0011573549,0.002166517,0.006093907,0.009129671,0.007473728,0.014277056,0.012512622,0.27207974],"study_design_scores_gemma":[0.00037515478,0.0011618146,0.6025436,0.0030732106,0.0038777082,0.005010192,0.015027859,0.11691786,0.038851257,0.0718642,0.1409606,0.00033644505],"about_ca_topic_score_codex":0.009021567,"about_ca_topic_score_gemma":0.00846807,"teacher_disagreement_score":0.009021567,"about_ca_system_score_codex":0.00088989764,"about_ca_system_score_gemma":0.0019713317,"threshold_uncertainty_score":0.03974116},"labels":[],"label_agreement":null},{"id":"W4394620374","doi":"10.7554/elife.94909","title":"Mining the neuroimaging literature","year":2024,"lang":"en","type":"article","venue":"eLife","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"National Institute of Mental Health; Canadian Institutes of Health Research; National Institutes of Health; Natural Sciences and Engineering Research Council of Canada; Fondation Brain Canada; Fonds de recherche du Québec; Canada First Research Excellence Fund; National Institute of Biomedical Imaging and Bioengineering; Chan Zuckerberg Initiative; Fonds de Recherche du Québec - Santé; Michael J. Fox Foundation for Parkinson's Research","keywords":"Neuroimaging; Psychology; Computer science; Neuroscience","score_opus":0.012355871193723202,"score_gpt":0.2787494241017696,"score_spread":0.2663935529080464,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394620374","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.017870316,0.041949373,0.12021897,0.008689008,0.0015608155,0.0019290546,0.70241135,0.01803083,0.08734033],"genre_scores_gemma":[0.04731111,0.028573858,0.36193332,0.0023491585,0.0010199604,0.0029995474,0.53603435,0.0027658003,0.017012974],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99625045,0.0007903059,0.0008982908,0.00086953636,0.0010546606,0.00013674569],"domain_scores_gemma":[0.98652124,0.0062770275,0.0020488198,0.001684455,0.0029264507,0.00054205157],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.005541783,0.0011345969,0.0007767133,0.04908979,0.0013884029,0.0038826705,0.0018373573,0.0012070317,0.02756043],"category_scores_gemma":[0.02740445,0.000624734,0.0018260378,0.028820388,0.00071727845,0.003656308,0.0040945066,0.0011270719,0.019538619],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001842286,0.00008267281,0.008138644,0.018760268,0.0005930754,0.001824314,0.0017510959,0.00056743785,0.0071208198,0.01879132,0.41508707,0.5270991],"study_design_scores_gemma":[0.000052520052,0.00003540134,0.009630018,0.00383202,0.0004944846,0.001245165,0.0008594461,0.0012584628,0.003497666,0.019361733,0.95967263,0.000060386024],"about_ca_topic_score_codex":0.004993643,"about_ca_topic_score_gemma":0.013760454,"teacher_disagreement_score":0.9944582,"about_ca_system_score_codex":0.0015651144,"about_ca_system_score_gemma":0.006301975,"threshold_uncertainty_score":0.09219885},"labels":[],"label_agreement":null},{"id":"W4394621145","doi":"10.7554/elife.94909.1","title":"Mining the neuroimaging literature","year":2024,"lang":"en","type":"preprint","venue":"eLife","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"","keywords":"Neuroimaging; Computer science; Data science; Psychology; Neuroscience","score_opus":0.01755183096316879,"score_gpt":0.29141702380512896,"score_spread":0.27386519284196015,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4394621145","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09589043,0.09575755,0.3176351,0.027683327,0.0030918207,0.0035706314,0.2829949,0.012464248,0.16091192],"genre_scores_gemma":[0.1746562,0.043815937,0.5540973,0.002844003,0.0017966785,0.0029785275,0.1996252,0.0021840183,0.018002186],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99510103,0.0012994575,0.0008640883,0.0009553319,0.0016125401,0.0001675337],"domain_scores_gemma":[0.974574,0.012815225,0.002810863,0.0021380002,0.0068550548,0.0008068494],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0062030554,0.0008787275,0.00085909583,0.050567303,0.0017022776,0.0039383187,0.0018466725,0.0009851472,0.011134255],"category_scores_gemma":[0.03548637,0.00045309044,0.0012251711,0.027118178,0.0010888102,0.0030927572,0.0036555212,0.0010374206,0.0054699723],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021572733,0.00011963918,0.014079348,0.0136640705,0.000491243,0.0038910117,0.0034300024,0.0016963753,0.0105866,0.03276995,0.17146966,0.7475864],"study_design_scores_gemma":[0.00006569803,0.000062855375,0.01557966,0.006259672,0.0006563212,0.0023126984,0.0029074638,0.0062698103,0.010731012,0.069427125,0.8856316,0.00009604064],"about_ca_topic_score_codex":0.005075876,"about_ca_topic_score_gemma":0.009723135,"teacher_disagreement_score":0.9494327,"about_ca_system_score_codex":0.0017437041,"about_ca_system_score_gemma":0.00711612,"threshold_uncertainty_score":0.037247837},"labels":[],"label_agreement":null},{"id":"W4395033543","doi":"10.3233/jpd-230305","title":"In Their Own Words: Fears Expressed by People with Parkinson’s Disease in an Online Symptom Database","year":2024,"lang":"en","type":"article","venue":"Journal of Parkinson s Disease","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; University Health Network","funders":"Parkinson's UK; Michael J. Fox Foundation for Parkinson's Research","keywords":"Parkinson's disease; Disease; Psychology; Database; Medicine; Computer science; Pathology","score_opus":0.016899460262921353,"score_gpt":0.2809668450857816,"score_spread":0.2640673848228603,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395033543","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.62196416,0.0040632044,0.021005591,0.010478217,0.00047617176,0.0007468111,0.3276829,0.00177771,0.011805253],"genre_scores_gemma":[0.80873305,0.0021131954,0.038605966,0.0018016081,0.00018911506,0.000580577,0.1450473,0.00017483113,0.0027543188],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99877363,0.00035056385,0.0003473935,0.00015862979,0.00029718335,0.00007260101],"domain_scores_gemma":[0.99249214,0.0052850554,0.0011211059,0.00039992505,0.00047687115,0.00022484973],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018698445,0.0003057177,0.00032492864,0.002417143,0.0005163131,0.0016884009,0.0004441938,0.00071389705,0.0043219686],"category_scores_gemma":[0.016849425,0.000078685865,0.00047330942,0.0026513657,0.00032321454,0.0022658166,0.0010842904,0.0004771536,0.00097428885],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015196612,0.00040277338,0.39681977,0.0082112625,0.00058458943,0.0020539453,0.018022658,0.0018800041,0.009257342,0.007854191,0.16319768,0.39019626],"study_design_scores_gemma":[0.00019208848,0.0004935563,0.54880613,0.0033410804,0.0008331976,0.0064681377,0.0368057,0.015530838,0.008511613,0.018408408,0.36029205,0.0003172575],"about_ca_topic_score_codex":0.004193129,"about_ca_topic_score_gemma":0.0074807014,"teacher_disagreement_score":0.0043219686,"about_ca_system_score_codex":0.00069930154,"about_ca_system_score_gemma":0.0007609889,"threshold_uncertainty_score":0.014458418},"labels":[],"label_agreement":null},{"id":"W4395166741","doi":"10.15468/dl.8hyy2d","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395166741","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006363287,0.000037112542,0.00006199594,0.000051700947,0.000012160979,0.0000073264414,0.9983475,0.000622825,0.00079575117],"genre_scores_gemma":[0.000170158,0.000032128388,0.00022234203,0.00004138536,0.0000027083015,0.00003664386,0.99891853,0.00013404727,0.00044203975],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99894446,0.0001378969,0.00013005883,0.0003564843,0.00026969486,0.00016143182],"domain_scores_gemma":[0.99781775,0.0005958588,0.00020206915,0.0005653701,0.0005139875,0.00030487913],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010571117,0.002416482,0.0016116281,0.005043335,0.0010511156,0.0025968084,0.003284166,0.0022759377,0.12542914],"category_scores_gemma":[0.0055185226,0.00094707654,0.0011923431,0.008983692,0.00052240206,0.0023781734,0.0026389847,0.0021678444,0.18285637],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035134904,0.000015404323,0.00034404555,0.00046970812,0.000014442539,0.000019512634,0.00002361164,0.00015412369,0.00012636647,0.00040099825,0.99675685,0.0016397795],"study_design_scores_gemma":[0.00008743609,0.000009296048,0.0018674935,0.00017055374,0.000014753743,0.00005744379,0.000077349956,0.0002929308,0.0002876573,0.0009195711,0.9961966,0.00001892764],"about_ca_topic_score_codex":0.023509623,"about_ca_topic_score_gemma":0.03808806,"teacher_disagreement_score":0.87457085,"about_ca_system_score_codex":0.001964179,"about_ca_system_score_gemma":0.002607491,"threshold_uncertainty_score":0.41960227},"labels":[],"label_agreement":null},{"id":"W4395187262","doi":"10.15468/dl.5sk797","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395187262","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000079238656,0.000038545208,0.00007040746,0.00005543627,0.0000136886465,0.000008560844,0.99823004,0.00069552596,0.0008085053],"genre_scores_gemma":[0.00017289337,0.00003201117,0.00024154532,0.000043134605,0.0000027281023,0.000040426643,0.9989655,0.00013085747,0.00037100102],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989428,0.00014056776,0.00013374163,0.0003577177,0.0002652442,0.00015999009],"domain_scores_gemma":[0.9979411,0.0005846332,0.0001901969,0.0005293122,0.00048091347,0.00027380884],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001028962,0.002522577,0.0015692366,0.004810698,0.0011202195,0.002552103,0.003193055,0.0024144517,0.09565625],"category_scores_gemma":[0.0051460764,0.00093607453,0.0012913678,0.008762821,0.00047888712,0.0022533927,0.0026113086,0.0022540435,0.16015106],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003917697,0.000018532488,0.00043012796,0.00051884,0.000016523947,0.000024223627,0.00002777324,0.00016822935,0.00016006375,0.0004553773,0.9964914,0.0016497944],"study_design_scores_gemma":[0.00009315953,0.000010229157,0.002005179,0.00017584793,0.000015232362,0.00005895496,0.00008188076,0.00032288465,0.00030748933,0.00093392323,0.9959745,0.000020643178],"about_ca_topic_score_codex":0.024319608,"about_ca_topic_score_gemma":0.040497832,"teacher_disagreement_score":0.9043437,"about_ca_system_score_codex":0.0019138085,"about_ca_system_score_gemma":0.0025322638,"threshold_uncertainty_score":0.32000202},"labels":[],"label_agreement":null},{"id":"W4395201578","doi":"10.15468/dl.2b4fmf","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395201578","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008674164,0.00005449133,0.000081217324,0.000059241404,0.000016009746,0.00000883885,0.9980286,0.00075930805,0.00090553204],"genre_scores_gemma":[0.00020678023,0.00004566047,0.00028948192,0.00005030643,0.0000031783454,0.00003635974,0.99880993,0.00012168614,0.00043668732],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99902,0.00012503573,0.00013038339,0.00034176806,0.00023750635,0.00014536774],"domain_scores_gemma":[0.9982256,0.00045095917,0.00014384225,0.00049066433,0.00047260273,0.0002162089],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00081445667,0.002210017,0.0015068601,0.0047596027,0.0010203886,0.002502374,0.003020635,0.0021674135,0.09453744],"category_scores_gemma":[0.0049449885,0.00085674354,0.0012380648,0.008876132,0.00043546426,0.0023266985,0.0023661442,0.0020474303,0.15380453],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039571667,0.000018976774,0.0004603943,0.0005651008,0.000017219018,0.000025483225,0.000023403307,0.00019548406,0.00016014594,0.00045113164,0.99580574,0.0022373884],"study_design_scores_gemma":[0.00007866835,0.000010081926,0.0017011208,0.00017006097,0.000013978051,0.00006400513,0.00008205405,0.00033314904,0.0002824087,0.00088995596,0.99635625,0.00001823318],"about_ca_topic_score_codex":0.019979462,"about_ca_topic_score_gemma":0.034083568,"teacher_disagreement_score":0.90546256,"about_ca_system_score_codex":0.0017472053,"about_ca_system_score_gemma":0.0022851469,"threshold_uncertainty_score":0.3162592},"labels":[],"label_agreement":null},{"id":"W4395215867","doi":"10.15468/dl.7yj9ze","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395215867","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00009554172,0.000044762764,0.00006307898,0.000050664505,0.000012264027,0.0000075555968,0.99838734,0.000553351,0.00078538625],"genre_scores_gemma":[0.00019714126,0.000035624293,0.00021136075,0.00003777456,0.0000025071333,0.000031623353,0.99901295,0.00009994034,0.00037106022],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989017,0.000143263,0.00014085614,0.0003739647,0.00027544596,0.00016478842],"domain_scores_gemma":[0.99802643,0.0005429384,0.00019107058,0.0005068941,0.0004641317,0.0002685118],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009580015,0.0023804654,0.0015353404,0.005503613,0.0010739906,0.0025126382,0.0029397449,0.0022777878,0.08867156],"category_scores_gemma":[0.004846757,0.0008743273,0.0011996558,0.009703834,0.0004675829,0.0022015802,0.0024777954,0.002020621,0.13983862],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046666413,0.000020494808,0.00057485583,0.0006122487,0.000020376396,0.000030465211,0.000029065674,0.00020789471,0.00018776242,0.00051328994,0.99580956,0.0019472698],"study_design_scores_gemma":[0.00008764065,0.000010182226,0.002245044,0.00017476428,0.000016853839,0.00007038391,0.0000839253,0.00030671267,0.00030630257,0.0008237503,0.9958549,0.000019544717],"about_ca_topic_score_codex":0.02467249,"about_ca_topic_score_gemma":0.0423986,"teacher_disagreement_score":0.91132843,"about_ca_system_score_codex":0.0018508838,"about_ca_system_score_gemma":0.0025900921,"threshold_uncertainty_score":0.29663593},"labels":[],"label_agreement":null},{"id":"W4395263243","doi":"10.15468/dl.6zk7b7","title":"Occurrence Download","year":2021,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014720005460750911,"score_gpt":0.23632839260909824,"score_spread":0.22160838714834732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395263243","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000066050045,0.000036944828,0.00006077809,0.000046104775,0.000010703495,0.0000070587207,0.99855953,0.00050142145,0.00071142294],"genre_scores_gemma":[0.00017269985,0.000035238612,0.00022253164,0.000036440444,0.0000023487714,0.000034841356,0.9989747,0.000114265036,0.00040687434],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988921,0.0001484435,0.00014704102,0.00037104366,0.00028029928,0.00016105601],"domain_scores_gemma":[0.99768925,0.0006685968,0.00021017858,0.00059204316,0.0005330193,0.00030689544],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010687612,0.0023054467,0.0016111188,0.005645703,0.001019059,0.0026247168,0.0031166198,0.002230645,0.10956248],"category_scores_gemma":[0.005682226,0.0009385509,0.0011953304,0.010166629,0.000513814,0.00222955,0.0025016314,0.002189718,0.15600312],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039340423,0.000017310242,0.00045372624,0.0005969909,0.000018564748,0.000025147274,0.000028247212,0.00019965401,0.00015813178,0.00050081004,0.99603504,0.0019271084],"study_design_scores_gemma":[0.00008017333,0.000008624214,0.0019555204,0.00017849036,0.000016031057,0.00005948512,0.000079475794,0.00027085093,0.00027796847,0.0008631015,0.99619234,0.000017955905],"about_ca_topic_score_codex":0.02376708,"about_ca_topic_score_gemma":0.03895926,"teacher_disagreement_score":0.89043754,"about_ca_system_score_codex":0.001974024,"about_ca_system_score_gemma":0.0027152167,"threshold_uncertainty_score":0.36652303},"labels":[],"label_agreement":null},{"id":"W4395343658","doi":"10.15468/dl.758zu8","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395343658","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000068489324,0.000040955394,0.0000657359,0.00004944985,0.000012595355,0.000007924752,0.99850607,0.0005209022,0.0007280132],"genre_scores_gemma":[0.00016667413,0.000035501456,0.00024570854,0.000041106527,0.0000024951096,0.000039857427,0.9989826,0.00011104532,0.00037511555],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989359,0.00014274636,0.0001440245,0.00035931106,0.000261926,0.00015606971],"domain_scores_gemma":[0.99792767,0.0005834234,0.00018870487,0.00052726356,0.000504268,0.00026871523],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010743145,0.0023683019,0.0016038496,0.0047553754,0.0011344823,0.0025118375,0.003188966,0.0023707107,0.10368836],"category_scores_gemma":[0.005503081,0.00091502373,0.0012527454,0.009043928,0.0005174162,0.0021627932,0.002662847,0.0021567405,0.16099589],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039651455,0.000016986049,0.00042945286,0.0005989608,0.000018124992,0.00002353247,0.000030068992,0.00018033465,0.00017753255,0.00050174334,0.99625325,0.0017304558],"study_design_scores_gemma":[0.00009230283,0.000009503002,0.0018765373,0.00019123018,0.000015939457,0.000056896442,0.00008473779,0.00026152478,0.00029224795,0.0008896206,0.99620974,0.000019741068],"about_ca_topic_score_codex":0.024563458,"about_ca_topic_score_gemma":0.043346092,"teacher_disagreement_score":0.89631164,"about_ca_system_score_codex":0.001981435,"about_ca_system_score_gemma":0.002752077,"threshold_uncertainty_score":0.3468721},"labels":[],"label_agreement":null},{"id":"W4395396930","doi":"10.15468/dl.7jt33y","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395396930","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007535234,0.00004451868,0.00007254993,0.000051048464,0.000013095235,0.00000764067,0.9983146,0.0006280243,0.0007932095],"genre_scores_gemma":[0.00017050082,0.000036993675,0.00023148605,0.000041181203,0.0000026104954,0.00003383605,0.9989876,0.00012394952,0.00037186293],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988213,0.00015378051,0.00015378548,0.000399317,0.00029936997,0.00017237586],"domain_scores_gemma":[0.9978508,0.00061962794,0.00019829096,0.0005474336,0.00049044215,0.00029348576],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011160467,0.0024659506,0.0016188818,0.0054930686,0.0011081304,0.0026789554,0.0030690609,0.0023220696,0.10630217],"category_scores_gemma":[0.005556387,0.00095590815,0.0012435184,0.009633352,0.000511283,0.0023027447,0.0026731726,0.0021661434,0.16421454],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042554257,0.000018396066,0.00048024466,0.00062191323,0.000019025398,0.000027386333,0.000029624935,0.00018526481,0.00018596476,0.0005266411,0.99596846,0.0018945101],"study_design_scores_gemma":[0.00008301894,0.000009390144,0.0019232943,0.00018791632,0.000015700418,0.00006282569,0.000079248886,0.00026834468,0.00029893446,0.0008676687,0.99618477,0.000018846971],"about_ca_topic_score_codex":0.022019282,"about_ca_topic_score_gemma":0.037293393,"teacher_disagreement_score":0.89369786,"about_ca_system_score_codex":0.0019071756,"about_ca_system_score_gemma":0.0027208521,"threshold_uncertainty_score":0.3556162},"labels":[],"label_agreement":null},{"id":"W4395407962","doi":"10.15468/dl.2zaryn","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395407962","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000714753,0.000040975665,0.00006578086,0.00004823051,0.000012598491,0.000007448032,0.9984156,0.00059588463,0.0007418925],"genre_scores_gemma":[0.00015899773,0.000034902758,0.00022223058,0.000039516643,0.0000025946472,0.000034714736,0.99903023,0.00012351183,0.00035333773],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988167,0.00015722582,0.00015503765,0.00040766815,0.00028897385,0.00017437006],"domain_scores_gemma":[0.99774987,0.00065068994,0.0002027942,0.0005855422,0.0005086163,0.00030250268],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011356198,0.002492097,0.0016483755,0.0055501433,0.0011250146,0.0026699633,0.0031506198,0.0023086974,0.106989406],"category_scores_gemma":[0.005657193,0.0009686231,0.0012865235,0.0096453475,0.000514803,0.0023139997,0.0026691838,0.0021936106,0.16846202],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004112952,0.00001821967,0.00045669838,0.00059158227,0.000018764693,0.00002591981,0.00002850089,0.00017494228,0.0001708842,0.0004714538,0.9962781,0.001723893],"study_design_scores_gemma":[0.00008929861,0.000009652945,0.0019663447,0.00018675132,0.0000165703,0.0000632972,0.0000814237,0.00027804883,0.000289391,0.0008964414,0.99610317,0.00001959775],"about_ca_topic_score_codex":0.021917976,"about_ca_topic_score_gemma":0.037348937,"teacher_disagreement_score":0.8930106,"about_ca_system_score_codex":0.0018686097,"about_ca_system_score_gemma":0.0026803405,"threshold_uncertainty_score":0.35791522},"labels":[],"label_agreement":null},{"id":"W4395408739","doi":"10.15468/dl.4tk5wj","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395408739","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007645024,0.000040070066,0.000071316426,0.000050844974,0.000013645303,0.000007814279,0.99832386,0.0006291075,0.00078686315],"genre_scores_gemma":[0.00016147319,0.00003257227,0.000236104,0.000038676615,0.0000025446818,0.000035432844,0.9990069,0.000120357014,0.00036588614],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99891496,0.00014272188,0.0001343392,0.00036172257,0.0002777069,0.00016862176],"domain_scores_gemma":[0.9979006,0.00058748946,0.00019352812,0.0005390901,0.00049646257,0.0002828041],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010368254,0.0024244552,0.0016231856,0.005077331,0.0011393153,0.0025367492,0.003165684,0.0023271623,0.10414033],"category_scores_gemma":[0.0051623653,0.00093236635,0.0012374765,0.009244063,0.00048714617,0.0021776834,0.0025884656,0.002244623,0.16506603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039130427,0.000018728666,0.00042952911,0.00052333844,0.000016463615,0.000024021976,0.000027096094,0.0001735289,0.00017409796,0.00045751792,0.99635124,0.001765283],"study_design_scores_gemma":[0.00009130382,0.0000105981635,0.0020583093,0.00017604543,0.00001586102,0.00006333285,0.00008360862,0.00029594544,0.000313236,0.0009490642,0.9959222,0.000020495954],"about_ca_topic_score_codex":0.022832748,"about_ca_topic_score_gemma":0.039095886,"teacher_disagreement_score":0.89585966,"about_ca_system_score_codex":0.0018812338,"about_ca_system_score_gemma":0.002615568,"threshold_uncertainty_score":0.34838408},"labels":[],"label_agreement":null},{"id":"W4395494250","doi":"10.15468/dl.2vhhva","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395494250","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008568948,0.000046430745,0.000062437466,0.000051282084,0.000013088177,0.000007711163,0.9984029,0.0005696748,0.00076077203],"genre_scores_gemma":[0.00018287654,0.00003574,0.00020671761,0.00003910822,0.000002506785,0.00003106608,0.9990589,0.00010126183,0.00034170802],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988562,0.00015028963,0.000145757,0.00039726993,0.00028028092,0.0001702144],"domain_scores_gemma":[0.99799913,0.0005390012,0.0001869778,0.0005245851,0.00046980535,0.00028043007],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010008881,0.0024155888,0.0016148539,0.0054682763,0.0011248072,0.0025896477,0.0030655842,0.002361146,0.09303879],"category_scores_gemma":[0.005163693,0.00088335975,0.0012231237,0.009885964,0.0004943044,0.0023028804,0.0026011951,0.0020847449,0.14363924],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044904766,0.000019499294,0.00048261933,0.0006054006,0.000019451283,0.000028001541,0.000027983388,0.0001908033,0.00017570853,0.0005000198,0.99614954,0.0017560459],"study_design_scores_gemma":[0.000092076094,0.000010094423,0.0020765406,0.0001815509,0.00001669285,0.00006907707,0.00008309272,0.00028406476,0.0002851635,0.00085372926,0.99602747,0.000020385285],"about_ca_topic_score_codex":0.026725484,"about_ca_topic_score_gemma":0.045680545,"teacher_disagreement_score":0.9069612,"about_ca_system_score_codex":0.0018577131,"about_ca_system_score_gemma":0.00271516,"threshold_uncertainty_score":0.3112458},"labels":[],"label_agreement":null},{"id":"W4395560999","doi":"10.15468/dl.2b7rnp","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395560999","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007583052,0.000040637566,0.000064939755,0.00004934904,0.000012906127,0.0000077258355,0.9984376,0.00055976614,0.00075118087],"genre_scores_gemma":[0.00016787114,0.00003384698,0.00022250558,0.000039261457,0.0000026840546,0.000034991135,0.99903,0.00011374806,0.00035511813],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988636,0.00015394726,0.00015154737,0.0003852458,0.00027647655,0.0001691775],"domain_scores_gemma":[0.99778616,0.0006453373,0.00020325193,0.0005639295,0.00050027727,0.00030097182],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011213832,0.0023573851,0.0015715342,0.0053596394,0.0010989499,0.0025977076,0.003039648,0.0022459957,0.103937134],"category_scores_gemma":[0.005523421,0.00090956886,0.0012289194,0.009249183,0.0005067308,0.0022372068,0.0025943937,0.0021589426,0.16018832],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004178807,0.000018791365,0.00049082004,0.00059921766,0.000018642428,0.000026953943,0.000029374254,0.00017970675,0.00018082248,0.00048572657,0.99615127,0.0017769097],"study_design_scores_gemma":[0.00008852661,0.000010045563,0.0021062659,0.00019283745,0.00001619486,0.0000653627,0.00008644127,0.0002754861,0.00028936984,0.00088423927,0.99596596,0.000019333324],"about_ca_topic_score_codex":0.020848062,"about_ca_topic_score_gemma":0.036602728,"teacher_disagreement_score":0.89606285,"about_ca_system_score_codex":0.001802147,"about_ca_system_score_gemma":0.0026142858,"threshold_uncertainty_score":0.34770435},"labels":[],"label_agreement":null},{"id":"W4395595955","doi":"10.15468/dl.257h92","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395595955","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006449417,0.00003647763,0.0000636095,0.000050142444,0.000011809462,0.0000073373135,0.99831223,0.00065549184,0.00079849386],"genre_scores_gemma":[0.0001739751,0.00003209986,0.00022722724,0.0000408059,0.0000026778457,0.000036093938,0.99891233,0.00014140502,0.00043339442],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989458,0.00013587122,0.00012986542,0.00035603865,0.00027022016,0.00016224325],"domain_scores_gemma":[0.9977933,0.00060195255,0.00020494737,0.00057877705,0.0005102767,0.00031079102],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010431088,0.002398089,0.0016118292,0.0051620174,0.0010531007,0.0025934118,0.003235408,0.002257235,0.12573048],"category_scores_gemma":[0.005519876,0.00094900327,0.0011983876,0.009111419,0.0005171702,0.0023777662,0.0026883087,0.0021623506,0.18157773],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003667907,0.000015868041,0.0003661382,0.00048700615,0.0000150249925,0.000020857968,0.000025057696,0.00015829007,0.00013593778,0.00041775036,0.99659866,0.0017228145],"study_design_scores_gemma":[0.000084222454,0.000009199226,0.001910737,0.00016992792,0.0000147606725,0.000057957826,0.000078369085,0.00028703973,0.00029073388,0.0009135354,0.9961647,0.000018877028],"about_ca_topic_score_codex":0.023155486,"about_ca_topic_score_gemma":0.03739977,"teacher_disagreement_score":0.8742695,"about_ca_system_score_codex":0.0019498662,"about_ca_system_score_gemma":0.0025863464,"threshold_uncertainty_score":0.42061037},"labels":[],"label_agreement":null},{"id":"W4395599148","doi":"10.15468/dl.297vsj","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395599148","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008703238,0.000047315887,0.000068249654,0.00005823734,0.000014555796,0.00000877746,0.9981664,0.0006729024,0.0008765139],"genre_scores_gemma":[0.00021358463,0.00004093143,0.00026019945,0.000046754074,0.0000032173577,0.000039678343,0.99885595,0.00013339793,0.00040633648],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99895144,0.00013317971,0.00013409529,0.0003530158,0.00026070772,0.00016752258],"domain_scores_gemma":[0.9978962,0.0005776511,0.0001805864,0.00055339327,0.0005289215,0.0002632377],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0008963016,0.002405433,0.0015263545,0.005214814,0.0011539932,0.0025431921,0.003061434,0.0023947845,0.1012432],"category_scores_gemma":[0.0055754236,0.0009231502,0.0012520219,0.009640852,0.00050125166,0.0024100945,0.002585719,0.0020908595,0.15823355],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039425533,0.0000183091,0.00046951263,0.0005733723,0.000016337368,0.000025757327,0.000027982229,0.00016521227,0.00016110402,0.00044037006,0.9963258,0.0017367098],"study_design_scores_gemma":[0.00009059336,0.000010519569,0.002048921,0.0001927806,0.000015561109,0.000066023385,0.00009873887,0.00029220537,0.00028871102,0.0008955795,0.99598,0.000020246951],"about_ca_topic_score_codex":0.025660709,"about_ca_topic_score_gemma":0.045536604,"teacher_disagreement_score":0.8987568,"about_ca_system_score_codex":0.0018850581,"about_ca_system_score_gemma":0.0026023793,"threshold_uncertainty_score":0.33869225},"labels":[],"label_agreement":null},{"id":"W4395657118","doi":"10.15468/dl.7w992j","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395657118","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006355165,0.00004357867,0.000058264955,0.000045004304,0.000011937197,0.000007116737,0.9985482,0.00052677747,0.00069564546],"genre_scores_gemma":[0.00016549267,0.00003892451,0.00020828305,0.000044365264,0.0000028675456,0.00003935582,0.99900913,0.00012424795,0.00036731144],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99881625,0.00015639579,0.00015723206,0.00041445915,0.00028149845,0.00017416124],"domain_scores_gemma":[0.9977557,0.00065318873,0.00021747682,0.0005727708,0.0005085533,0.00029228444],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010882273,0.002459571,0.001659534,0.0055525443,0.0011103473,0.0027259507,0.0031358895,0.0023874932,0.11256413],"category_scores_gemma":[0.005818219,0.00095563306,0.0012404243,0.009867954,0.0005078294,0.0023166058,0.0027573444,0.0021538404,0.17632209],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004159377,0.000016426722,0.00040554916,0.0006646099,0.000019803661,0.000025130885,0.000025076492,0.00015608236,0.00017442893,0.00045018623,0.99640405,0.0016171527],"study_design_scores_gemma":[0.0000933081,0.000009891367,0.0017949025,0.00019963007,0.000018237499,0.00006280794,0.0000687397,0.00022245297,0.00028947118,0.00086236617,0.99635834,0.000019841145],"about_ca_topic_score_codex":0.019601064,"about_ca_topic_score_gemma":0.034460276,"teacher_disagreement_score":0.88743585,"about_ca_system_score_codex":0.0017805403,"about_ca_system_score_gemma":0.0025928013,"threshold_uncertainty_score":0.37656456},"labels":[],"label_agreement":null},{"id":"W4395658422","doi":"10.15468/dl.3wu8yq","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395658422","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000064614935,0.000037263326,0.000069936534,0.000049142203,0.000012155419,0.000007886597,0.99837947,0.0006146535,0.00076477876],"genre_scores_gemma":[0.00016490195,0.000034179902,0.00024430244,0.000041247942,0.0000025767013,0.000038019214,0.99896765,0.00013569575,0.0003714888],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99885845,0.00015136908,0.00015008343,0.00038243475,0.00028897516,0.0001686917],"domain_scores_gemma":[0.99773115,0.0006478458,0.00020490197,0.00059681904,0.00051896984,0.00030031762],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011294854,0.0024804482,0.0016429386,0.0052106013,0.0011496098,0.002694597,0.0031880687,0.0023954555,0.111466564],"category_scores_gemma":[0.0055501014,0.00096754864,0.0013169021,0.009216783,0.0005108555,0.002335175,0.002727896,0.0022495622,0.16979545],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039813953,0.000017509212,0.00042442547,0.00057561375,0.000017910488,0.000025756075,0.000030092613,0.00016996417,0.00016971426,0.00049327745,0.99631447,0.0017214398],"study_design_scores_gemma":[0.0000881015,0.000009165811,0.0018379502,0.00017853685,0.000015766065,0.000062343686,0.00007914822,0.0002637826,0.0002956525,0.000940811,0.99620867,0.000020050622],"about_ca_topic_score_codex":0.023127656,"about_ca_topic_score_gemma":0.038615063,"teacher_disagreement_score":0.8885334,"about_ca_system_score_codex":0.0019728285,"about_ca_system_score_gemma":0.0027161152,"threshold_uncertainty_score":0.3728928},"labels":[],"label_agreement":null},{"id":"W4395734363","doi":"10.15468/dl.amu555","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395734363","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007996439,0.000046334117,0.000059202703,0.000050734656,0.000012864385,0.00000767283,0.99852717,0.00050689,0.0007092768],"genre_scores_gemma":[0.00019060032,0.00004010393,0.00020910763,0.000043947868,0.0000029669325,0.000035977395,0.9990074,0.00010019154,0.0003697896],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989312,0.00014296177,0.00013970815,0.0003709664,0.00025508186,0.00016002136],"domain_scores_gemma":[0.9979836,0.0005865049,0.00019753833,0.0005010202,0.00045820719,0.0002731825],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009742461,0.002386511,0.0015893241,0.005658948,0.0010475197,0.0024593135,0.0028975604,0.0023334092,0.10318698],"category_scores_gemma":[0.005178293,0.000883835,0.0012145455,0.009912598,0.00047843202,0.002304055,0.0024747446,0.0020327112,0.1561344],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044274646,0.000019785275,0.00049868703,0.00067450514,0.000019829251,0.000028378394,0.000026963226,0.00019980571,0.0001845377,0.0004586416,0.9958711,0.0019734444],"study_design_scores_gemma":[0.0000940763,0.000011045509,0.0021058018,0.00019956984,0.00001765559,0.00006712849,0.00008001684,0.0002860626,0.00029559896,0.0008638095,0.99595886,0.000020357424],"about_ca_topic_score_codex":0.0197645,"about_ca_topic_score_gemma":0.032799702,"teacher_disagreement_score":0.89681304,"about_ca_system_score_codex":0.0017445028,"about_ca_system_score_gemma":0.0023903057,"threshold_uncertainty_score":0.34519482},"labels":[],"label_agreement":null},{"id":"W4395850718","doi":"10.15468/dl.947jeu","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395850718","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000072718016,0.00004006594,0.00005992761,0.000050559538,0.00001232118,0.0000074239892,0.99849236,0.0005176472,0.00074697],"genre_scores_gemma":[0.00019034614,0.000036286572,0.00021390845,0.000042536252,0.0000027330934,0.000036922,0.99896896,0.00011661998,0.00039176125],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998922,0.00014730282,0.00014044305,0.00036963954,0.0002589993,0.00016167469],"domain_scores_gemma":[0.99783665,0.00064404984,0.00020257097,0.00054892234,0.000473944,0.00029392977],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00107281,0.002270229,0.0015184828,0.0054440238,0.0010359214,0.0025750569,0.003030744,0.0022219317,0.10532497],"category_scores_gemma":[0.005543865,0.00090674846,0.0012041344,0.009531453,0.0004982299,0.0022557096,0.0024809416,0.0020873316,0.15197262],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041996416,0.000017995619,0.00046134708,0.0006043999,0.00001825772,0.000025521062,0.000027154829,0.00018891749,0.0001653234,0.0005042811,0.996136,0.0018086851],"study_design_scores_gemma":[0.00008727335,0.000009871251,0.0019381368,0.00017894832,0.00001597901,0.00006217149,0.00007769326,0.00028699322,0.00028012993,0.00091898633,0.9961249,0.000018794291],"about_ca_topic_score_codex":0.021275876,"about_ca_topic_score_gemma":0.034526695,"teacher_disagreement_score":0.894675,"about_ca_system_score_codex":0.0018692538,"about_ca_system_score_gemma":0.0025622293,"threshold_uncertainty_score":0.35234714},"labels":[],"label_agreement":null},{"id":"W4395862434","doi":"10.15468/dl.9bhamp","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395862434","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007734106,0.00004493012,0.00005247,0.000044116885,0.000013027416,0.000007712424,0.9984825,0.0005031256,0.00077469024],"genre_scores_gemma":[0.00017092003,0.00003863793,0.000192777,0.000043869946,0.000002881291,0.000031003277,0.999038,0.0001087658,0.00037321722],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989052,0.00014060539,0.00015117563,0.00036681956,0.00026115426,0.0001750353],"domain_scores_gemma":[0.9979235,0.0005202454,0.00019940134,0.000550841,0.00051582407,0.00029016516],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009810321,0.002296986,0.0015713501,0.0058281324,0.001026149,0.0025473733,0.002782068,0.0021215212,0.10510872],"category_scores_gemma":[0.0049589123,0.0008886187,0.0012467448,0.009870227,0.00049080426,0.002220128,0.0026813133,0.002005554,0.1674391],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047181635,0.000018615185,0.00047623823,0.0006740328,0.000020131098,0.000027199916,0.000027774635,0.00015489927,0.00021824849,0.0004784689,0.9959495,0.0019077818],"study_design_scores_gemma":[0.000080126,0.000010361804,0.001965339,0.00018341953,0.000015904207,0.00006219816,0.00007544756,0.00020048903,0.000285687,0.00075264025,0.99635077,0.000017635186],"about_ca_topic_score_codex":0.020618407,"about_ca_topic_score_gemma":0.037050493,"teacher_disagreement_score":0.89489126,"about_ca_system_score_codex":0.001785224,"about_ca_system_score_gemma":0.002546506,"threshold_uncertainty_score":0.3516237},"labels":[],"label_agreement":null},{"id":"W4395888176","doi":"10.15468/dl.acbhwj","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395888176","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000082766695,0.000049599646,0.00006354094,0.00004839964,0.000013259782,0.000008406459,0.9984981,0.00052524515,0.00071058277],"genre_scores_gemma":[0.00017776777,0.000039742557,0.0002130816,0.000041646545,0.0000027711815,0.00003760748,0.99905556,0.00010250498,0.0003293465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989178,0.00014449585,0.00014654166,0.00037742683,0.00025120977,0.000162422],"domain_scores_gemma":[0.99796474,0.0005825649,0.00019714562,0.00051362504,0.00046767926,0.0002742779],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010369894,0.0024767711,0.0016973093,0.0054459963,0.0011015099,0.00255626,0.0031678786,0.00244509,0.09765249],"category_scores_gemma":[0.005245946,0.00091595977,0.0012513328,0.009778652,0.00050665694,0.0022797207,0.0026383481,0.002113894,0.15408066],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000049739454,0.000021759504,0.00048237742,0.0007470997,0.00002145421,0.000031623378,0.00002918602,0.0001876768,0.00022075298,0.00048438803,0.99582255,0.0019014821],"study_design_scores_gemma":[0.000103233746,0.000012049518,0.0021150792,0.00021857754,0.000019005753,0.0000730879,0.000085599866,0.0002619903,0.0003234007,0.0008637418,0.9959027,0.000021397253],"about_ca_topic_score_codex":0.020731764,"about_ca_topic_score_gemma":0.0346796,"teacher_disagreement_score":0.9023475,"about_ca_system_score_codex":0.0017454912,"about_ca_system_score_gemma":0.0024993145,"threshold_uncertainty_score":0.32668012},"labels":[],"label_agreement":null},{"id":"W4395897363","doi":"10.15468/dl.b7tfgg","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395897363","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007333823,0.00004156556,0.00006571922,0.000049294937,0.000012593878,0.00000799109,0.99837744,0.0005827647,0.00078925287],"genre_scores_gemma":[0.00016345095,0.000034354038,0.0002255251,0.000039450402,0.0000026264136,0.000036754147,0.9990206,0.00012041146,0.00035680278],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99879324,0.00016488307,0.00015762809,0.00040357813,0.0003033135,0.0001773755],"domain_scores_gemma":[0.9976587,0.00067584636,0.00020704698,0.0006007546,0.0005498487,0.00030777676],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001186834,0.0024650348,0.0016584774,0.005558221,0.0011864078,0.002738249,0.0031766393,0.0024128407,0.10566157],"category_scores_gemma":[0.0057935137,0.0009769226,0.0012884984,0.009782152,0.00053001445,0.002332379,0.0026951532,0.0022334082,0.16647796],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042008047,0.000019407262,0.00046674276,0.0006252228,0.000018657305,0.0000271808,0.000030160521,0.00017074702,0.00018566771,0.0004904501,0.99619436,0.0017295212],"study_design_scores_gemma":[0.00008893874,0.000009717534,0.001999793,0.00018883412,0.000015728758,0.00006325355,0.00008531689,0.00025078445,0.00029522053,0.0008387191,0.9961438,0.000020017258],"about_ca_topic_score_codex":0.022688895,"about_ca_topic_score_gemma":0.039458342,"teacher_disagreement_score":0.8943384,"about_ca_system_score_codex":0.0019389413,"about_ca_system_score_gemma":0.002790615,"threshold_uncertainty_score":0.3534732},"labels":[],"label_agreement":null},{"id":"W4395905565","doi":"10.15468/dl.a5jpjg","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395905565","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007362957,0.000039161358,0.000073799696,0.00005150025,0.000012935035,0.000007807355,0.99830437,0.00064174447,0.00079507614],"genre_scores_gemma":[0.00017055048,0.00003542852,0.00025387894,0.00004208987,0.000002613685,0.000039232822,0.9989213,0.00014438147,0.00039059337],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989649,0.00013758906,0.00013518492,0.0003503696,0.00025726846,0.00015479751],"domain_scores_gemma":[0.9979394,0.0005833515,0.00019124952,0.0005312542,0.00048190478,0.00027276535],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010271574,0.0023854738,0.0015749579,0.0051448797,0.0010702127,0.0024987902,0.0030931658,0.0022384631,0.10392651],"category_scores_gemma":[0.005359764,0.00094037474,0.0012410159,0.009434066,0.0004922994,0.0022653413,0.0026637157,0.0021940616,0.16119477],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038419035,0.000016609882,0.00041452493,0.00055428053,0.000016899823,0.00002384669,0.00002912374,0.00017580402,0.00016337333,0.00048653365,0.996327,0.0017536305],"study_design_scores_gemma":[0.0000857413,0.0000089488,0.001877811,0.00017968645,0.000014832291,0.000058496815,0.000084086685,0.00026951838,0.00027725237,0.0008976319,0.9962268,0.00001924889],"about_ca_topic_score_codex":0.024346523,"about_ca_topic_score_gemma":0.04157057,"teacher_disagreement_score":0.89607346,"about_ca_system_score_codex":0.00187773,"about_ca_system_score_gemma":0.0025367336,"threshold_uncertainty_score":0.34766883},"labels":[],"label_agreement":null},{"id":"W4395944032","doi":"10.15468/dl.ccvk7w","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395944032","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000074404015,0.000041662344,0.00006591913,0.000048235233,0.000012133664,0.000007606338,0.9984048,0.0005410502,0.00080414565],"genre_scores_gemma":[0.00018018472,0.000036421145,0.00022739306,0.000041656738,0.000002636181,0.00003754314,0.99895704,0.00012376948,0.00039336266],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988971,0.00015168851,0.00014341707,0.00037146168,0.00027198004,0.00016436764],"domain_scores_gemma":[0.99785703,0.0006365227,0.00019190715,0.0005415512,0.00049406284,0.00027898583],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010904927,0.0022834637,0.0015694085,0.005420659,0.0011085116,0.0026376923,0.0030612855,0.002216898,0.115459785],"category_scores_gemma":[0.0054891547,0.00093163014,0.0012022698,0.009684314,0.0004875726,0.0022075954,0.0025619345,0.0020984416,0.16632126],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003973645,0.000017793267,0.00045202684,0.0005874089,0.000017902406,0.000024945408,0.000027772778,0.00016955857,0.00016390042,0.00048175035,0.9962619,0.001755364],"study_design_scores_gemma":[0.000084597,0.000009076534,0.0019152355,0.00018859543,0.00001592898,0.000060469127,0.0000818963,0.0002440569,0.00027500596,0.00086545653,0.9962405,0.000019092893],"about_ca_topic_score_codex":0.021218337,"about_ca_topic_score_gemma":0.036944967,"teacher_disagreement_score":0.8845402,"about_ca_system_score_codex":0.0017677416,"about_ca_system_score_gemma":0.0025737137,"threshold_uncertainty_score":0.38625145},"labels":[],"label_agreement":null},{"id":"W4395983813","doi":"10.15468/dl.fhu7rf","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4395983813","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000093965034,0.00004541521,0.00007771083,0.000057893158,0.00001565491,0.000010059554,0.99811065,0.0007484091,0.00084027246],"genre_scores_gemma":[0.00018662505,0.00003635472,0.0002836867,0.0000454545,0.0000028571021,0.00004378547,0.99888545,0.00012952586,0.00038636543],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989685,0.00013879093,0.00013664218,0.00034979143,0.00025471114,0.00015162966],"domain_scores_gemma":[0.9979748,0.0005697158,0.00017948559,0.00054104364,0.0004727102,0.00026228264],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010479749,0.0024768102,0.0015597088,0.0050697397,0.001149646,0.0024906797,0.0031432526,0.0023922215,0.09544045],"category_scores_gemma":[0.0052518067,0.0009068867,0.0013577647,0.008886496,0.0005097089,0.0022152965,0.0025675504,0.002147222,0.15441853],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000043808865,0.000021177108,0.00049444,0.0005916826,0.000018845156,0.00002923464,0.000031029693,0.00020615023,0.00018401949,0.00045522678,0.99594635,0.0019780467],"study_design_scores_gemma":[0.0000995474,0.00001167888,0.0021240576,0.00018685065,0.000016930764,0.00007042547,0.0000926165,0.00034954638,0.00032543513,0.0009551579,0.9957457,0.000022038203],"about_ca_topic_score_codex":0.023460489,"about_ca_topic_score_gemma":0.039895803,"teacher_disagreement_score":0.90455955,"about_ca_system_score_codex":0.0018827548,"about_ca_system_score_gemma":0.0025650226,"threshold_uncertainty_score":0.3192801},"labels":[],"label_agreement":null},{"id":"W4396000464","doi":"10.15468/dl.e4mj37","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396000464","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006770032,0.000041240884,0.000055482764,0.000043890523,0.0000120367085,0.000007299456,0.998531,0.0004995257,0.00074178504],"genre_scores_gemma":[0.00017215816,0.000040061903,0.00021593638,0.00004569431,0.0000030161177,0.00003671963,0.9989594,0.00011652582,0.00041047164],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989213,0.00014773096,0.00014868048,0.0003770538,0.00024504573,0.00016020208],"domain_scores_gemma":[0.99792576,0.0005684336,0.00020833613,0.00052508607,0.00049033546,0.0002820374],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009691382,0.0024603794,0.0015847477,0.0054509756,0.00097372977,0.0024902842,0.002977854,0.0021007892,0.10953482],"category_scores_gemma":[0.00528572,0.00092297787,0.0012218974,0.009636605,0.00046877738,0.0022500267,0.0025494837,0.0019325395,0.16338593],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000433448,0.000016874128,0.00043085098,0.0006345217,0.000018941846,0.000024012541,0.000023200444,0.00016566203,0.00016163477,0.00043930393,0.9962591,0.001782678],"study_design_scores_gemma":[0.00009076412,0.000010747226,0.0017508556,0.00017616914,0.000016683853,0.00005986666,0.000064118736,0.00024120814,0.00026736225,0.00088725064,0.99641573,0.000019316838],"about_ca_topic_score_codex":0.01891578,"about_ca_topic_score_gemma":0.033210423,"teacher_disagreement_score":0.8904652,"about_ca_system_score_codex":0.0016340464,"about_ca_system_score_gemma":0.0024871281,"threshold_uncertainty_score":0.36643052},"labels":[],"label_agreement":null},{"id":"W4396011879","doi":"10.15468/dl.dty3bj","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396011879","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007437296,0.000045180695,0.00007294852,0.000056476627,0.000014126982,0.000008050279,0.9982033,0.0007025828,0.0008230434],"genre_scores_gemma":[0.00017117306,0.00003673433,0.00024013902,0.000044161818,0.0000027450196,0.00003631301,0.9989448,0.00013601058,0.00038788238],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988943,0.00014524096,0.00013717658,0.00037626663,0.00027899153,0.00016809792],"domain_scores_gemma":[0.9978841,0.0005957078,0.00019524652,0.00054127106,0.0004949933,0.0002886897],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00104055,0.0024921873,0.0015993684,0.005258694,0.0011387098,0.0026499461,0.0032005932,0.0023991389,0.10226793],"category_scores_gemma":[0.0054421937,0.00092373625,0.0012722998,0.0090873595,0.00049528613,0.0023253583,0.0027391072,0.0022499089,0.16396873],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037971826,0.000017360511,0.00042381298,0.0005714659,0.000017534257,0.000024997038,0.000027494227,0.00016795115,0.00016832091,0.00047184667,0.99632084,0.0017504666],"study_design_scores_gemma":[0.00008044391,0.000009384897,0.001767169,0.00017957622,0.000015371039,0.00006091952,0.00007790826,0.00027572716,0.00028648728,0.0008741269,0.99635375,0.00001910656],"about_ca_topic_score_codex":0.023637697,"about_ca_topic_score_gemma":0.04056375,"teacher_disagreement_score":0.8977321,"about_ca_system_score_codex":0.0019009834,"about_ca_system_score_gemma":0.0026695405,"threshold_uncertainty_score":0.3421203},"labels":[],"label_agreement":null},{"id":"W4396047689","doi":"10.15468/dl.cu2jp3","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396047689","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007376887,0.00004045922,0.000069785834,0.00004749296,0.000013090229,0.0000078114435,0.9983614,0.0005960523,0.0007901017],"genre_scores_gemma":[0.00016775668,0.000035399982,0.00024736646,0.00004291805,0.000002530848,0.000038164315,0.9989526,0.00012663977,0.0003865147],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990036,0.00012628538,0.00012884026,0.00034540787,0.00024375897,0.00015205231],"domain_scores_gemma":[0.9980889,0.0005152475,0.00018173695,0.0005099166,0.00044840528,0.0002557801],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009717576,0.002323148,0.0015896684,0.005008763,0.0010932768,0.00245984,0.0031057203,0.0023066623,0.10262642],"category_scores_gemma":[0.004998083,0.00093073724,0.0012602528,0.008904993,0.00048591627,0.002223901,0.0026154872,0.0021531754,0.16592573],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038479644,0.00001695348,0.0004489634,0.00059345516,0.00001755779,0.000024556522,0.000029877461,0.00017996863,0.00018590187,0.0004854521,0.9961255,0.0018532595],"study_design_scores_gemma":[0.00007809927,0.000008722509,0.0018888314,0.0001837297,0.00001505695,0.000058019577,0.00008129687,0.00025950995,0.0002781948,0.0008165543,0.9963135,0.0000185313],"about_ca_topic_score_codex":0.022789322,"about_ca_topic_score_gemma":0.040887628,"teacher_disagreement_score":0.89737356,"about_ca_system_score_codex":0.0018781724,"about_ca_system_score_gemma":0.0024772484,"threshold_uncertainty_score":0.34331954},"labels":[],"label_agreement":null},{"id":"W4396063432","doi":"10.15468/dl.fk27zp","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; Internet privacy; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396063432","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00009224597,0.000049676153,0.000060285023,0.000048289086,0.0000134166485,0.000007999278,0.9984877,0.0005093841,0.000731007],"genre_scores_gemma":[0.00020566738,0.000040297775,0.00021059102,0.000040625033,0.0000029347818,0.000034912253,0.99903905,0.00009151107,0.00033449108],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989465,0.00014069033,0.00014005005,0.00036647386,0.0002497436,0.00015662187],"domain_scores_gemma":[0.9980823,0.0005597177,0.00018746317,0.00048265088,0.0004234842,0.00026433746],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00096490287,0.0023334934,0.0015857379,0.0055246125,0.0010550527,0.0024811516,0.002890322,0.002290525,0.09339168],"category_scores_gemma":[0.0050232797,0.00086272083,0.0012088474,0.009543084,0.00048015555,0.0021852374,0.0024438575,0.001995908,0.14113517],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000050901417,0.000022106276,0.00058295095,0.00075311115,0.000022482855,0.000033761768,0.00002899241,0.00022014664,0.0002131235,0.00050449726,0.99551195,0.0020560953],"study_design_scores_gemma":[0.00009935114,0.000011963144,0.0022932105,0.00020824169,0.000018780685,0.000076601114,0.00008399639,0.00030186764,0.00030663627,0.0008669847,0.9957113,0.000020925288],"about_ca_topic_score_codex":0.019527633,"about_ca_topic_score_gemma":0.033635918,"teacher_disagreement_score":0.90660834,"about_ca_system_score_codex":0.0016773815,"about_ca_system_score_gemma":0.002393552,"threshold_uncertainty_score":0.31242633},"labels":[],"label_agreement":null},{"id":"W4396148782","doi":"10.15468/dl.f2dg25","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396148782","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00006935087,0.000041099298,0.00007264118,0.000050670387,0.000012697568,0.00000792433,0.9983804,0.00057264423,0.0007925593],"genre_scores_gemma":[0.00016860524,0.000036678946,0.00024042474,0.000042779753,0.0000026182865,0.00003836318,0.9989723,0.00012620486,0.00037201136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988416,0.00015997511,0.00015500661,0.00039302846,0.0002815359,0.00016895175],"domain_scores_gemma":[0.99774075,0.0006813681,0.00020137851,0.00056772534,0.000512315,0.00029650875],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011627466,0.0023237318,0.0015977778,0.0054526487,0.0011212128,0.0027006343,0.0031564555,0.0022921818,0.110884316],"category_scores_gemma":[0.0059700296,0.0009496535,0.0012472359,0.009423867,0.0005069074,0.00234649,0.0026461403,0.0022097605,0.16313708],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040263596,0.00001765113,0.00044587848,0.00061355793,0.000018743141,0.000025821419,0.000027936923,0.00018096507,0.00016226509,0.0005059402,0.99616325,0.0017978013],"study_design_scores_gemma":[0.00008464769,0.000009047135,0.001809151,0.00019455818,0.000016266558,0.000060981132,0.000079300815,0.00026310945,0.00027117808,0.00090825453,0.9962843,0.000019123892],"about_ca_topic_score_codex":0.020843875,"about_ca_topic_score_gemma":0.036002655,"teacher_disagreement_score":0.8891157,"about_ca_system_score_codex":0.0018521502,"about_ca_system_score_gemma":0.0026431696,"threshold_uncertainty_score":0.37094498},"labels":[],"label_agreement":null},{"id":"W4396188219","doi":"10.15468/dl.e2nvef","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396188219","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000632838,0.000037736732,0.000062209816,0.000051914987,0.000012299872,0.000007301158,0.9983429,0.00062563794,0.0007966176],"genre_scores_gemma":[0.00017037604,0.00003273061,0.00022305794,0.00004208951,0.0000027433532,0.000036568417,0.99891555,0.00013541541,0.00044142545],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99894685,0.00013755796,0.00012980824,0.0003566436,0.0002677208,0.0001614242],"domain_scores_gemma":[0.9978307,0.0005930927,0.00020048031,0.00056062033,0.0005122784,0.00030288516],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001052933,0.002408206,0.0016098284,0.0050134943,0.0010530755,0.002605458,0.0032686747,0.0022714895,0.12558761],"category_scores_gemma":[0.0055189487,0.0009438891,0.001196375,0.008934614,0.0005208848,0.0023821546,0.0026444294,0.0021545086,0.18357539],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035128007,0.0000152242355,0.00034388487,0.00047539343,0.000014568425,0.000019508525,0.0000236408,0.00015372928,0.00012678234,0.00040256567,0.9967494,0.0016402575],"study_design_scores_gemma":[0.00008706223,0.000009259374,0.0018517952,0.00017177651,0.000014805569,0.00005712142,0.0000770561,0.00028993524,0.00028456718,0.0009240646,0.9962136,0.000018899933],"about_ca_topic_score_codex":0.023492409,"about_ca_topic_score_gemma":0.03814739,"teacher_disagreement_score":0.8744124,"about_ca_system_score_codex":0.0019556354,"about_ca_system_score_gemma":0.0026015267,"threshold_uncertainty_score":0.4201324},"labels":[],"label_agreement":null},{"id":"W4396256051","doi":"10.15468/dl.gpjymw","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396256051","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000655828,0.000036044043,0.000063295236,0.000050730294,0.000011818143,0.000007649062,0.9983375,0.0006505608,0.0007767584],"genre_scores_gemma":[0.00016916056,0.000031548774,0.0002207162,0.00003884131,0.000002667441,0.000036803227,0.9989379,0.00013459378,0.000427884],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99889165,0.00014516136,0.00013811376,0.00037432343,0.00028364596,0.00016714045],"domain_scores_gemma":[0.99763227,0.00065055233,0.00021671498,0.0006288876,0.0005439331,0.0003276261],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010977167,0.0024277696,0.0016368881,0.005344376,0.0010657208,0.0026443703,0.0032852385,0.002284988,0.12159464],"category_scores_gemma":[0.0057354462,0.00094692403,0.0012241598,0.009532767,0.0005356628,0.0023921586,0.0027132484,0.0022521247,0.17887262],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037376994,0.000017140135,0.000374603,0.0004883157,0.000015287847,0.000021405216,0.000026409236,0.0001617236,0.00013788593,0.00042883249,0.99657935,0.0017117504],"study_design_scores_gemma":[0.00008360377,0.000009240491,0.0018982927,0.00016602306,0.000014461281,0.000055581542,0.00008065439,0.0002825109,0.00028614982,0.00091447594,0.9961903,0.000018667337],"about_ca_topic_score_codex":0.023592146,"about_ca_topic_score_gemma":0.037808716,"teacher_disagreement_score":0.87840533,"about_ca_system_score_codex":0.0019923656,"about_ca_system_score_gemma":0.0027007845,"threshold_uncertainty_score":0.40677458},"labels":[],"label_agreement":null},{"id":"W4396298415","doi":"10.15468/dl.h5au7g","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396298415","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007905841,0.000044413646,0.000060286133,0.00005332515,0.000014146909,0.000008370942,0.99851567,0.00051587913,0.00070892816],"genre_scores_gemma":[0.0001855577,0.00003764685,0.00021814034,0.000045811033,0.000003029136,0.000037734662,0.9990005,0.0000988543,0.00037262886],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988167,0.00016321482,0.00015709273,0.00041633693,0.00027377662,0.0001729217],"domain_scores_gemma":[0.9977106,0.00066098705,0.00021306523,0.000607768,0.000511325,0.00029634743],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010765981,0.002392449,0.0016177779,0.0053997813,0.0010994043,0.0025305937,0.0031371075,0.0024199593,0.10165283],"category_scores_gemma":[0.0056812586,0.0009058702,0.001270778,0.009456033,0.0005040685,0.0023603768,0.0025781463,0.0021039594,0.15045092],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000045541325,0.000020286388,0.00047186748,0.00064944674,0.000019212703,0.000026626109,0.000025147454,0.00019756156,0.00015818057,0.0004584719,0.9961372,0.0017903716],"study_design_scores_gemma":[0.000101736056,0.000012278494,0.002051908,0.00020383579,0.000017333743,0.00006712765,0.000081057005,0.00029760867,0.00028106172,0.0009411718,0.99592364,0.000021237178],"about_ca_topic_score_codex":0.020546189,"about_ca_topic_score_gemma":0.035821762,"teacher_disagreement_score":0.89834714,"about_ca_system_score_codex":0.0017568776,"about_ca_system_score_gemma":0.00251904,"threshold_uncertainty_score":0.34006262},"labels":[],"label_agreement":null},{"id":"W4396308825","doi":"10.15468/dl.jywat9","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396308825","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000073034964,0.00004131689,0.00007322856,0.000054429853,0.000014315139,0.000008241412,0.9982174,0.00066297373,0.0008550753],"genre_scores_gemma":[0.0001674389,0.00003599889,0.0002491734,0.000042476004,0.0000025647253,0.000036940728,0.9989201,0.00012902831,0.00041633454],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988764,0.0001477866,0.00014696718,0.00038103547,0.00027879688,0.00016906919],"domain_scores_gemma":[0.9978027,0.0005773838,0.00019204589,0.00059928326,0.0005431499,0.00028544624],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010216185,0.002280364,0.0015702336,0.005141379,0.0011265835,0.0025971925,0.0032195677,0.0021798634,0.0983851],"category_scores_gemma":[0.005513772,0.0009098544,0.001250304,0.0094861975,0.0005122547,0.0023462668,0.0025893457,0.0021365536,0.15821974],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003649132,0.000016454675,0.00043731203,0.0005164351,0.000016042979,0.000023283039,0.000027965449,0.00018028641,0.00014486501,0.0004926014,0.99631065,0.0017976714],"study_design_scores_gemma":[0.0000706867,0.0000085314,0.0017452865,0.00016351322,0.000013608615,0.000058119287,0.000082769846,0.0002695219,0.00026048053,0.00089813303,0.9964109,0.00001840896],"about_ca_topic_score_codex":0.028580481,"about_ca_topic_score_gemma":0.04911732,"teacher_disagreement_score":0.9016149,"about_ca_system_score_codex":0.0020253544,"about_ca_system_score_gemma":0.002848964,"threshold_uncertainty_score":0.32913095},"labels":[],"label_agreement":null},{"id":"W4396311853","doi":"10.15468/dl.jqxns8","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396311853","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000084469306,0.00004462568,0.000066765366,0.000056043544,0.0000146303055,0.000008440692,0.99835384,0.00059338956,0.0007778064],"genre_scores_gemma":[0.00019222072,0.000036587335,0.00024465978,0.000045561606,0.0000030749818,0.000041542393,0.998934,0.000113785725,0.0003886081],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998928,0.00014505899,0.00013501849,0.00037119945,0.00025829146,0.00016232827],"domain_scores_gemma":[0.9978769,0.00062559644,0.00020109891,0.00053900015,0.00047188104,0.0002854615],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010289488,0.0024517565,0.0015582659,0.00521427,0.0011067636,0.0025018002,0.0030952017,0.0024327915,0.10146931],"category_scores_gemma":[0.0053319396,0.0009023061,0.0012452233,0.0090012085,0.00050438184,0.0022291732,0.0026001846,0.002166965,0.15614536],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039086583,0.00001915081,0.0004734816,0.00058318285,0.000017635219,0.000025530093,0.000027087191,0.00018741244,0.00016520033,0.00043013395,0.996232,0.001800106],"study_design_scores_gemma":[0.00009534472,0.000011529205,0.0020792955,0.00019492772,0.00001701519,0.00006818693,0.00008896009,0.00032044618,0.00030338252,0.00092567,0.9958747,0.000020576015],"about_ca_topic_score_codex":0.020444069,"about_ca_topic_score_gemma":0.035808448,"teacher_disagreement_score":0.8985307,"about_ca_system_score_codex":0.0018278995,"about_ca_system_score_gemma":0.0024795602,"threshold_uncertainty_score":0.33944863},"labels":[],"label_agreement":null},{"id":"W4396378460","doi":"10.15468/dl.j66v83","title":"Occurrence Download","year":2024,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.013410855924799302,"score_gpt":0.24025687544637875,"score_spread":0.22684601952157946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396378460","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000084253385,0.00004202908,0.00005463302,0.00004537914,0.000011116416,0.0000074469413,0.99853027,0.00047239405,0.0007525205],"genre_scores_gemma":[0.00018782627,0.000037300157,0.00020800925,0.000040006886,0.0000025306526,0.00003216637,0.99902797,0.00009170882,0.00037252292],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99901414,0.00013252605,0.00013674502,0.0003206334,0.00023886834,0.0001571133],"domain_scores_gemma":[0.9980282,0.0005154155,0.00020605502,0.0004940373,0.00048198536,0.00027433145],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00091334525,0.0021189707,0.0013964688,0.0056401487,0.0009792555,0.0023401058,0.0026984122,0.0020516093,0.09175154],"category_scores_gemma":[0.0047782916,0.0007909415,0.0011108409,0.009567416,0.00046271915,0.0020511544,0.0024354395,0.0018775208,0.13914499],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048227663,0.000019207218,0.00059112965,0.0007011592,0.000021097214,0.00003153859,0.000030460493,0.00019342887,0.0002262189,0.00054859906,0.99541134,0.0021775237],"study_design_scores_gemma":[0.000079869824,0.000009897739,0.0022213652,0.00019701855,0.000016367256,0.0000641334,0.000078818506,0.00023392182,0.00030425418,0.00083071523,0.9959453,0.000018365154],"about_ca_topic_score_codex":0.024367124,"about_ca_topic_score_gemma":0.04379055,"teacher_disagreement_score":0.9082485,"about_ca_system_score_codex":0.0017818996,"about_ca_system_score_gemma":0.0025995912,"threshold_uncertainty_score":0.30693948},"labels":[],"label_agreement":null},{"id":"W4396476091","doi":"10.15468/dl.jqrz42","title":"Occurrence Download","year":2020,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.016410385392551713,"score_gpt":0.23380037428754785,"score_spread":0.21738998889499614,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396476091","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000063996915,0.000037284917,0.00006205391,0.000051740695,0.000012281666,0.000007333186,0.99834836,0.00062244106,0.0007945688],"genre_scores_gemma":[0.00017105871,0.000032261454,0.00022376701,0.000041656316,0.0000027403462,0.00003654304,0.99891746,0.00013405224,0.00044056584],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99893755,0.00013850049,0.00013147123,0.00035931353,0.00027103352,0.00016216023],"domain_scores_gemma":[0.99779785,0.0006004403,0.00020373463,0.00057053566,0.00051907945,0.0003083579],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010587283,0.0024073264,0.0016088285,0.005051011,0.0010554338,0.0025983318,0.003272287,0.0022741584,0.12485197],"category_scores_gemma":[0.005530287,0.0009453771,0.0011981758,0.009013552,0.00052261667,0.0023709813,0.0026448132,0.0021651005,0.18198334],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035313915,0.00001546125,0.00034953962,0.00047378882,0.000014578845,0.000019703122,0.000023718954,0.00015523903,0.00012707878,0.000400994,0.9967359,0.0016486356],"study_design_scores_gemma":[0.00008744127,0.000009344466,0.0018832019,0.0001702369,0.000014835692,0.000057600257,0.00007732799,0.00029178016,0.000286239,0.0009157386,0.99618727,0.000018922752],"about_ca_topic_score_codex":0.023497894,"about_ca_topic_score_gemma":0.03815606,"teacher_disagreement_score":0.87514806,"about_ca_system_score_codex":0.0019597495,"about_ca_system_score_gemma":0.0026085926,"threshold_uncertainty_score":0.41767144},"labels":[],"label_agreement":null},{"id":"W4396490674","doi":"10.15468/dl.jt4qcd","title":"Occurrence Download","year":2023,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Computer science; World Wide Web","score_opus":0.018161949851733437,"score_gpt":0.2455588629157748,"score_spread":0.22739691306404136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396490674","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000079640675,0.00004255441,0.000065201355,0.000050210732,0.000013436595,0.000007675226,0.9984018,0.0005805179,0.00075896824],"genre_scores_gemma":[0.00017960441,0.00003584358,0.00023518104,0.000042370655,0.0000026582518,0.000036662695,0.9989641,0.000118727,0.0003848232],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989349,0.00013638723,0.00013918283,0.0003629697,0.00026577484,0.00016070233],"domain_scores_gemma":[0.99804866,0.0005200028,0.00018508219,0.0005166031,0.00046987264,0.00025979112],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009758445,0.0023756507,0.0016034532,0.0051240814,0.0010871358,0.0024716246,0.0030680206,0.002348861,0.099617034],"category_scores_gemma":[0.0050212094,0.00090730935,0.0012877334,0.009309001,0.0004987854,0.0022034706,0.0025981907,0.0021188357,0.16254967],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041013845,0.000018293804,0.00047976756,0.00056738243,0.000017781662,0.00002479196,0.000028608967,0.00018527964,0.0001815654,0.00046101256,0.9961611,0.0018334247],"study_design_scores_gemma":[0.000089040805,0.000010124822,0.0021090002,0.0001847761,0.00001597757,0.00006368632,0.000086726366,0.00028650276,0.00029568252,0.0008359303,0.99600285,0.000019622295],"about_ca_topic_score_codex":0.023683876,"about_ca_topic_score_gemma":0.03972195,"teacher_disagreement_score":0.900383,"about_ca_system_score_codex":0.001906505,"about_ca_system_score_gemma":0.0025455712,"threshold_uncertainty_score":0.3332522},"labels":[],"label_agreement":null},{"id":"W4396509049","doi":"10.5737/23688076342151","title":"Decision support for breast cancer screening decisions: A single case pre-/post-test study","year":2024,"lang":"en","type":"article","venue":"Canadian Oncology Nursing Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"University of Ottawa","funders":"","keywords":"Breast cancer; Test (biology); Cancer; Breast cancer screening; Medicine; Gynecology; Psychology; Mammography; Internal medicine; Biology","score_opus":0.045224041524912814,"score_gpt":0.3694838525167961,"score_spread":0.3242598109918833,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396509049","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99763846,0.00006221646,0.00019707196,0.00013702555,0.000023627948,0.000901752,0.0002019608,0.000011608969,0.0008262789],"genre_scores_gemma":[0.993468,0.00010603626,0.0012693129,0.0002477196,0.00004510514,0.0031076737,0.0005269904,0.00001084802,0.0012182936],"study_design_codex":"observational","study_design_gemma":"case_report","domain_scores_codex":[0.99203473,0.003501803,0.00081433874,0.0009147269,0.0017685004,0.00096594973],"domain_scores_gemma":[0.9531488,0.021888085,0.007621648,0.003975909,0.007684063,0.0056815487],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01022282,0.0006099653,0.0013261782,0.0015134253,0.0040618964,0.0017994273,0.0012504542,0.0021615361,0.004846357],"category_scores_gemma":[0.048775077,0.0009617777,0.0015241541,0.00085146213,0.0014653257,0.0021702524,0.0015378761,0.0035491788,0.001536757],"study_design_candidate":"case_report","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.021188347,0.21705446,0.6036746,0.0008820039,0.0009407176,0.005192864,0.052354064,0.00088157755,0.0015625918,0.0007893981,0.0061304676,0.089349024],"study_design_scores_gemma":[0.00503308,0.16046056,0.77174675,0.0006343497,0.0006648888,0.0029934242,0.037088837,0.003409026,0.0037105838,0.0018158337,0.011890152,0.0005523907],"about_ca_topic_score_codex":0.0068208724,"about_ca_topic_score_gemma":0.011033658,"teacher_disagreement_score":0.01022282,"about_ca_system_score_codex":0.0036351937,"about_ca_system_score_gemma":0.0035533498,"threshold_uncertainty_score":0.054064095},"labels":[],"label_agreement":null},{"id":"W4396891504","doi":"10.2196/55090","title":"Research on Traditional Chinese Medicine: Domain Knowledge Graph Completion and Quality Evaluation","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Knowledge graph; Quality (philosophy); Data science; Medicine; Information retrieval","score_opus":0.16631549335061777,"score_gpt":0.48018129451457126,"score_spread":0.31386580116395346,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4396891504","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.39697525,0.013693393,0.5617771,0.0035320076,0.0002093263,0.0017246889,0.0021451493,0.0023343298,0.017608754],"genre_scores_gemma":[0.7260808,0.0034543404,0.26562735,0.00031766717,0.00007207718,0.00047637266,0.002287364,0.0001930696,0.0014910254],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9861968,0.0065344605,0.00076123,0.001555417,0.004735191,0.00021677515],"domain_scores_gemma":[0.91954905,0.05191972,0.0074456325,0.0059879916,0.01394407,0.0011535872],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014825783,0.0010873056,0.000920293,0.009317312,0.00085309084,0.003054083,0.0016630469,0.00086075265,0.0029046033],"category_scores_gemma":[0.07998782,0.00038841105,0.0013796162,0.007958561,0.0014005962,0.005681299,0.0018457589,0.0010388928,0.0004025065],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000510544,0.0005154137,0.046168715,0.0033981036,0.00072451134,0.00017471048,0.0022516325,0.043392893,0.0041379556,0.021915382,0.005747843,0.87106234],"study_design_scores_gemma":[0.00023343031,0.0020388933,0.105768174,0.0018646654,0.0011869923,0.0009748217,0.004571462,0.7141392,0.027462598,0.09573895,0.045678016,0.00034276125],"about_ca_topic_score_codex":0.012424201,"about_ca_topic_score_gemma":0.009852777,"teacher_disagreement_score":0.014825783,"about_ca_system_score_codex":0.0037044536,"about_ca_system_score_gemma":0.0030944587,"threshold_uncertainty_score":0.07840711},"labels":[],"label_agreement":null},{"id":"W4398291336","doi":"10.7910/dvn/e4tlzv/hohftg","title":"AddFile2_Character list_131ch.pdf","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Character (mathematics); Divergence (linguistics); Bayesian probability; Evolutionary biology; Phylogenetics; Biology; Character evolution; Computer science; Artificial intelligence; Mathematics; Clade; Genetics","score_opus":0.016186787784778948,"score_gpt":0.2500295851177992,"score_spread":0.23384279733302027,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398291336","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000047459354,0.0000106097805,0.000044414148,0.000029492701,0.000011496375,0.000008338155,0.99883264,0.00032701207,0.000688587],"genre_scores_gemma":[0.00027646602,0.0000197625,0.00029966346,0.000060173574,0.000007637476,0.000066310284,0.9977041,0.00024082379,0.0013251362],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992611,0.000075725984,0.00007179486,0.00027286782,0.00016006229,0.0001583992],"domain_scores_gemma":[0.99743867,0.0008201599,0.00023195473,0.0005056115,0.0006489822,0.00035460418],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0009911096,0.002445446,0.0014577972,0.0040831477,0.0012420872,0.0028264686,0.0028085741,0.0017763095,0.4351558],"category_scores_gemma":[0.0047148503,0.0010574935,0.0011019349,0.005791454,0.0005711447,0.002550323,0.0025560558,0.0016757055,0.3183727],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004641362,0.000016477821,0.0004533947,0.0004647712,0.000011247372,0.000016078577,0.000017351975,0.00007203843,0.00014176911,0.0002750408,0.9969531,0.0015322348],"study_design_scores_gemma":[0.0002112215,0.0000142886565,0.0026364948,0.0001785109,0.000022090831,0.00005868481,0.00007633647,0.00017922238,0.0005740873,0.0010036499,0.9950147,0.000030711435],"about_ca_topic_score_codex":0.017111145,"about_ca_topic_score_gemma":0.03034195,"teacher_disagreement_score":0.5648442,"about_ca_system_score_codex":0.0018505924,"about_ca_system_score_gemma":0.0019958017,"threshold_uncertainty_score":0.8056817},"labels":[],"label_agreement":null},{"id":"W4398332633","doi":"10.7910/dvn/e4tlzv/jppiux","title":"AddFile5_SupplTables.rar","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Divergence (linguistics); Bayesian probability; Phylogenetics; Evolutionary biology; Molecular clock; Biology; Econometrics; Statistics; Mathematics; Genetics; Gene","score_opus":0.015591472539328423,"score_gpt":0.2474706816469962,"score_spread":0.2318792091076678,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398332633","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00004629668,0.000027526205,0.000114452465,0.000031719985,0.000015421861,0.000007731264,0.99821717,0.000717491,0.0008221554],"genre_scores_gemma":[0.00046242287,0.00004732339,0.00065268046,0.00010231087,0.000010904139,0.000086997956,0.9968167,0.0006186016,0.0012020302],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999005,0.00014014175,0.00011138945,0.00038943026,0.0001939871,0.00016009531],"domain_scores_gemma":[0.9966563,0.0013973744,0.00030477004,0.00075286854,0.00055340613,0.0003352857],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0016462135,0.0023683063,0.0018293713,0.0035868746,0.00096952287,0.0032668544,0.0033393977,0.0018728006,0.4036287],"category_scores_gemma":[0.007336681,0.0013157524,0.0014071675,0.004867043,0.00047007683,0.0019248181,0.0024822354,0.0019606282,0.30243695],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047515554,0.000015844993,0.0005789298,0.0007935124,0.00003296021,0.000014865855,0.00001942397,0.00023648843,0.00016614875,0.00065499,0.9948768,0.0025626759],"study_design_scores_gemma":[0.00027883603,0.000014766611,0.0026201615,0.00026954184,0.00003594244,0.000060879436,0.000049028597,0.00037384083,0.0006041883,0.0032102312,0.9924454,0.00003718902],"about_ca_topic_score_codex":0.011750457,"about_ca_topic_score_gemma":0.023590945,"teacher_disagreement_score":0.5963713,"about_ca_system_score_codex":0.0019952937,"about_ca_system_score_gemma":0.0024230902,"threshold_uncertainty_score":0.85065126},"labels":[],"label_agreement":null},{"id":"W4398362790","doi":"10.7910/dvn/bxiy5w/wjpfdk","title":"GRAY_2017.RData","year":2019,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Vector Institute; Princess Margaret Cancer Centre","funders":"","keywords":"Gray (unit); Computer science; Medicine; Nuclear medicine","score_opus":0.024273683221745945,"score_gpt":0.2766527202934732,"score_spread":0.25237903707172726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398362790","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00019766574,0.0000594003,0.00012261588,0.00007387842,0.000033923374,0.000017448507,0.99670714,0.0016332016,0.0011547308],"genre_scores_gemma":[0.0005007759,0.0000688314,0.00041529143,0.00007510299,0.0000069768644,0.000053575648,0.9977448,0.00025917735,0.00087537075],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990994,0.000103065475,0.00010142116,0.00029213596,0.00025302666,0.00015094482],"domain_scores_gemma":[0.9987012,0.00033886524,0.00010487919,0.00039521547,0.00026669083,0.000193061],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009862267,0.0034261295,0.0014683343,0.0039986805,0.00092533993,0.0023319998,0.0028131532,0.0021248206,0.10798643],"category_scores_gemma":[0.0036976403,0.0011592292,0.0017477917,0.0043427963,0.00061973644,0.0015577103,0.0025860167,0.0019652385,0.1244223],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010383101,0.000037144662,0.0005423318,0.0005533512,0.000033016768,0.000026663465,0.000024347755,0.0003713579,0.00032015066,0.0006396355,0.9942913,0.0030568675],"study_design_scores_gemma":[0.00040924113,0.000025764512,0.0018610252,0.00016169238,0.000039550232,0.00006604672,0.00005181526,0.00092239917,0.0011436046,0.0016284294,0.99365604,0.00003451762],"about_ca_topic_score_codex":0.019242631,"about_ca_topic_score_gemma":0.03230043,"teacher_disagreement_score":0.89201355,"about_ca_system_score_codex":0.0017535206,"about_ca_system_score_gemma":0.002648404,"threshold_uncertainty_score":0.36125058},"labels":[],"label_agreement":null},{"id":"W4398500242","doi":"10.7910/dvn/8yvinu/wsc0ke","title":"MedRed_dev.txt","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada)","funders":"","keywords":"Computer science","score_opus":0.018083430317117368,"score_gpt":0.25796225420834357,"score_spread":0.2398788238912262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398500242","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003087006,0.000038663435,0.00016287247,0.00011177303,0.000057292003,0.000020683248,0.9948348,0.002578888,0.0018862897],"genre_scores_gemma":[0.0006879052,0.0000256958,0.00053201494,0.000080902355,0.000014150242,0.000065598,0.9970636,0.00038718607,0.0011430315],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99837226,0.0002750504,0.00015330467,0.00049038784,0.00044590302,0.000263036],"domain_scores_gemma":[0.9955031,0.0017095228,0.0003020965,0.0012130692,0.000841839,0.00043029172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017936417,0.0034623062,0.0011921738,0.0036880143,0.0011878862,0.0027174847,0.0035004187,0.00225812,0.117342606],"category_scores_gemma":[0.0076221153,0.0009288347,0.0016968005,0.0036717083,0.0006875323,0.001816453,0.0023033177,0.0021669401,0.1452296],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006967527,0.000023665181,0.0005270105,0.0002205485,0.000013338962,0.000019900333,0.000011914884,0.00019726598,0.00008958665,0.00028116675,0.99683005,0.0017159415],"study_design_scores_gemma":[0.00048473943,0.000056205183,0.0037741933,0.00013866492,0.000030358615,0.0001311439,0.00007442775,0.001694311,0.0017040411,0.0019441312,0.98992443,0.000043360687],"about_ca_topic_score_codex":0.015054389,"about_ca_topic_score_gemma":0.03490561,"teacher_disagreement_score":0.117342606,"about_ca_system_score_codex":0.0019954075,"about_ca_system_score_gemma":0.0022886905,"threshold_uncertainty_score":0.3925501},"labels":[],"label_agreement":null},{"id":"W4398585633","doi":"10.7910/dvn/ktwyz6/d3jq3c","title":"risk_time.png","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"HEC Montréal","funders":"","keywords":"Computer science","score_opus":0.016870334338275183,"score_gpt":0.25610710631564504,"score_spread":0.23923677197736987,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398585633","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00014030609,0.00008133639,0.00011155707,0.00015737982,0.000028549988,0.000011778444,0.99647945,0.0014822194,0.0015075257],"genre_scores_gemma":[0.0012562982,0.0001414978,0.00045323127,0.00019176905,0.00003532907,0.000075607626,0.99552274,0.0006722861,0.0016513445],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990144,0.00015326729,0.00010612192,0.00031820018,0.00024839616,0.00015961689],"domain_scores_gemma":[0.9973348,0.00082793116,0.00029892527,0.0007889216,0.0003613552,0.00038816643],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011963648,0.00263529,0.0017625961,0.0048673805,0.00073230657,0.0038235933,0.003423755,0.0027286746,0.2810079],"category_scores_gemma":[0.008221388,0.0010091936,0.0013283449,0.007044847,0.0005206897,0.002128013,0.0028642986,0.0015074578,0.24487242],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006919927,0.00001938946,0.00063151686,0.00043013232,0.000029883535,0.000021170137,0.000013834416,0.0002465453,0.000060701517,0.000598858,0.9947643,0.003114475],"study_design_scores_gemma":[0.00040501193,0.00003299916,0.0020970912,0.0002674698,0.00003480181,0.00010831757,0.00004703138,0.0008956446,0.00048169383,0.0032438517,0.9923509,0.000035130302],"about_ca_topic_score_codex":0.010790043,"about_ca_topic_score_gemma":0.01486603,"teacher_disagreement_score":0.7189921,"about_ca_system_score_codex":0.0015982403,"about_ca_system_score_gemma":0.0015765032,"threshold_uncertainty_score":0.940065},"labels":[],"label_agreement":null},{"id":"W4398616996","doi":"10.7910/dvn/e4tlzv/uzdcte","title":"AddFile6_InputFiles&amp;OutputTrees.rar","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Divergence (linguistics); Bayesian probability; Phylogenetics; Evolutionary biology; Molecular clock; Biology; Econometrics; Computer science; Artificial intelligence; Mathematics; Genetics; Gene","score_opus":0.02553551625509093,"score_gpt":0.26428219775165757,"score_spread":0.23874668149656664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398616996","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008323753,0.000018723507,0.00018969242,0.00005132235,0.000019075467,0.000017527118,0.9964149,0.0021214073,0.0010840277],"genre_scores_gemma":[0.00046298263,0.00003339146,0.0009766151,0.00006932947,0.000010463432,0.0001431237,0.9958532,0.00126519,0.0011857289],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99891067,0.00015652427,0.00011930699,0.000388169,0.00024499168,0.00018035618],"domain_scores_gemma":[0.9965229,0.0012010533,0.00020882065,0.000901833,0.0008075858,0.00035786844],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0018994559,0.002965903,0.001567856,0.0037545073,0.0010526318,0.003206157,0.0033443402,0.0015543315,0.2999017],"category_scores_gemma":[0.008363574,0.0014267763,0.0016915237,0.0048916386,0.0005850103,0.0022512602,0.002821613,0.001855156,0.28030396],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006916959,0.00001729888,0.00047485754,0.00050718995,0.000019371091,0.000013962298,0.000024330438,0.00016483276,0.00014942171,0.00047648954,0.99577636,0.0023067284],"study_design_scores_gemma":[0.00028923567,0.000017475062,0.0020624988,0.00019312762,0.000026287567,0.00004590365,0.00006899779,0.0005529844,0.00097023736,0.002799541,0.9929327,0.000040942105],"about_ca_topic_score_codex":0.0096657695,"about_ca_topic_score_gemma":0.016295604,"teacher_disagreement_score":0.7000983,"about_ca_system_score_codex":0.0017955353,"about_ca_system_score_gemma":0.002212644,"threshold_uncertainty_score":0.99860525},"labels":[],"label_agreement":null},{"id":"W4398701832","doi":"10.7910/dvn/pl2xfd/39iv6k","title":"cses_coding and merging.do","year":2017,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Coding (social sciences); Computer science; Computer graphics (images); Mathematics; Statistics","score_opus":0.020898824348153027,"score_gpt":0.28903525155229837,"score_spread":0.2681364272041453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398701832","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00038278764,0.000038822927,0.0007004183,0.00008659118,0.00003056598,0.00004775912,0.99196666,0.005643868,0.0011025046],"genre_scores_gemma":[0.00078625785,0.00003765481,0.001861271,0.00008166596,0.000008372361,0.00019972827,0.9952554,0.0011572572,0.00061231863],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99890065,0.00012675568,0.00015988726,0.00042754193,0.00021832497,0.00016679321],"domain_scores_gemma":[0.9971539,0.0010160395,0.0002075652,0.000751637,0.0005558113,0.0003150464],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0014746863,0.0025372668,0.0010783626,0.0057146433,0.0008435513,0.0021312635,0.0028963843,0.0014853991,0.15515482],"category_scores_gemma":[0.0075698476,0.0009994317,0.0017798558,0.0063103894,0.00065945624,0.0021404792,0.002846118,0.0017961629,0.12074544],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009490904,0.000029058883,0.0010955471,0.00040333782,0.0000320651,0.000023406925,0.00003500634,0.00037657784,0.00026216023,0.00069253985,0.99305844,0.0038969358],"study_design_scores_gemma":[0.00054114027,0.000026622662,0.0036702785,0.00019689635,0.000042119857,0.000105859515,0.00009664944,0.002051618,0.0022775254,0.0043591224,0.98657644,0.0000557212],"about_ca_topic_score_codex":0.013285828,"about_ca_topic_score_gemma":0.026269702,"teacher_disagreement_score":0.8448452,"about_ca_system_score_codex":0.0017099419,"about_ca_system_score_gemma":0.0031214724,"threshold_uncertainty_score":0.51904464},"labels":[],"label_agreement":null},{"id":"W4398715339","doi":"10.7910/dvn/8yvinu/ch7jee","title":"MedRed_AMT_labels.csv","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bell (Canada)","funders":"","keywords":"Computer science","score_opus":0.01783809722862782,"score_gpt":0.25566350298115953,"score_spread":0.23782540575253172,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4398715339","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00039275954,0.00006698951,0.00016467315,0.00015606963,0.00008399612,0.000029285151,0.99451524,0.0024551407,0.0021358968],"genre_scores_gemma":[0.0011243908,0.00006502155,0.00089396583,0.00014955575,0.00003543791,0.000100373,0.99451256,0.00037791897,0.0027407603],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988965,0.00016412007,0.00010607955,0.00034274737,0.00028681522,0.00020368933],"domain_scores_gemma":[0.9970897,0.0008793483,0.00026183142,0.00085458637,0.00056430785,0.0003501517],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012203186,0.0025924991,0.0011637688,0.0042492007,0.001150615,0.0023076872,0.002210948,0.0026411314,0.18318655],"category_scores_gemma":[0.0061971266,0.00086245,0.0010491487,0.004244287,0.00063154375,0.00156494,0.0022544488,0.0016407118,0.23256324],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006058139,0.000022184586,0.00034037797,0.000335899,0.000008665927,0.000021006526,0.000018595016,0.000103289334,0.0001568683,0.00024404941,0.99618906,0.0024994197],"study_design_scores_gemma":[0.00025878396,0.000042422795,0.003355161,0.0001473691,0.000018609782,0.0001281,0.00007078444,0.0012956843,0.0011910658,0.0014010962,0.99205583,0.00003491535],"about_ca_topic_score_codex":0.013122545,"about_ca_topic_score_gemma":0.030029872,"teacher_disagreement_score":0.81681347,"about_ca_system_score_codex":0.0017882583,"about_ca_system_score_gemma":0.0019686953,"threshold_uncertainty_score":0.61282},"labels":[],"label_agreement":null},{"id":"W4399039386","doi":"10.1109/access.2024.3405529","title":"Fully Automated Scholarly Search for Biomedical Systematic Literature Reviews","year":2024,"lang":"en","type":"article","venue":"IEEE Access","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Toronto Metropolitan University","funders":"National Research Council Canada","keywords":"Computer science; Information retrieval; Benchmark (surveying); Suite; Set (abstract data type); Generative grammar; Precision and recall; Process (computing); Recall; Field (mathematics); Data mining; Data science; Artificial intelligence","score_opus":0.05230406408299433,"score_gpt":0.3899084845378799,"score_spread":0.3376044204548856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399039386","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07805236,0.01899574,0.76588315,0.008411957,0.0003309525,0.00556161,0.057894807,0.05658002,0.008289301],"genre_scores_gemma":[0.20364918,0.002421201,0.7498225,0.0010565349,0.00015568332,0.001699697,0.039053496,0.00077165134,0.0013700498],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9820696,0.009872389,0.0026794288,0.0017026095,0.0033878423,0.00028814073],"domain_scores_gemma":[0.9063554,0.07210947,0.0055035753,0.009017754,0.006087315,0.0009265273],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.019146098,0.0013833736,0.0018876267,0.016435884,0.0012942263,0.003539823,0.002324832,0.0022590894,0.006500077],"category_scores_gemma":[0.12486797,0.0012099147,0.0030739068,0.010122711,0.00068706984,0.005453538,0.0045472835,0.0014152303,0.0039347005],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013591605,0.00048175626,0.012148473,0.01626461,0.0009940645,0.001029506,0.002223436,0.03309813,0.014894661,0.022385899,0.094549984,0.80057037],"study_design_scores_gemma":[0.0015374057,0.00086374313,0.008306161,0.0020704584,0.0013029022,0.0022146844,0.0016414546,0.73862547,0.025657266,0.09815013,0.11931577,0.0003145236],"about_ca_topic_score_codex":0.004804221,"about_ca_topic_score_gemma":0.013529074,"teacher_disagreement_score":0.9808539,"about_ca_system_score_codex":0.002093167,"about_ca_system_score_gemma":0.011420039,"threshold_uncertainty_score":0.10125548},"labels":[],"label_agreement":null},{"id":"W4399074020","doi":"10.1101/2024.05.27.24307990","title":"The Neurodegenerative Disease Knowledge Portal: Propelling Discovery Through the Sharing of Neurodegenerative Disease Genomic Resources","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"","keywords":"Disease; Neuroscience; Medicine; Biology; Pathology","score_opus":0.029938643568661445,"score_gpt":0.29155718309422246,"score_spread":0.261618539525561,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399074020","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021662876,0.004708947,0.71472013,0.033171926,0.0015141506,0.0017423239,0.045191575,0.13249344,0.044794567],"genre_scores_gemma":[0.076686904,0.005147225,0.78388435,0.005609091,0.00086420064,0.001199996,0.109588616,0.010738108,0.0062814886],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98827165,0.0051486786,0.0014829608,0.0017908943,0.0025459456,0.000759871],"domain_scores_gemma":[0.94947106,0.02052397,0.0018753326,0.016704107,0.0037427342,0.0076827607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028711943,0.0013755438,0.0018233089,0.008453482,0.002058285,0.012290976,0.005660966,0.0034703508,0.011553955],"category_scores_gemma":[0.053665303,0.0014211364,0.0021345406,0.010470524,0.002080256,0.017684976,0.034533292,0.005081131,0.010511847],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020921063,0.00062494655,0.013714547,0.0038741403,0.0006133402,0.0039978293,0.0063263085,0.0071345754,0.010837061,0.12546754,0.35483772,0.47047988],"study_design_scores_gemma":[0.0006528922,0.00024739152,0.004342366,0.0013265107,0.00020116336,0.0011301889,0.002447893,0.020453703,0.0057697133,0.24531965,0.71775675,0.00035182963],"about_ca_topic_score_codex":0.0042547425,"about_ca_topic_score_gemma":0.005061804,"teacher_disagreement_score":0.028711943,"about_ca_system_score_codex":0.0014654456,"about_ca_system_score_gemma":0.0068197907,"threshold_uncertainty_score":0.1518451},"labels":[],"label_agreement":null},{"id":"W4399097616","doi":"10.31234/osf.io/c83d7","title":"Cognitive Mode Detectable with Task-Based fMRI: Language (LAN)","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Provincial Health Services Authority; University of British Columbia","funders":"","keywords":"Default mode network; Cognition; Neuroscience; Cognitive science; Mode (computer interface); Psychology; Computer science; Cognitive psychology; Anatomy; Medicine; Human–computer interaction","score_opus":0.011314670507106994,"score_gpt":0.29068742619561977,"score_spread":0.2793727556885128,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399097616","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93230385,0.00039411464,0.053548943,0.0003215273,0.000048016518,0.00016161389,0.0006466941,0.00052560924,0.012049613],"genre_scores_gemma":[0.9823111,0.00010560873,0.01541354,0.00023413835,0.000041614887,0.00014255295,0.0002889718,0.00012966974,0.0013327862],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.9996939,0.00006116929,0.000024054596,0.00010724257,0.00005251603,0.00006116509],"domain_scores_gemma":[0.998749,0.00054643763,0.00030360918,0.00020976212,0.000094573705,0.00009659972],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008619242,0.0005258637,0.00022662962,0.00053344975,0.0002110323,0.00079101237,0.00037623596,0.0005842519,0.0026876004],"category_scores_gemma":[0.003358292,0.00017725045,0.00023687408,0.00029274196,0.0007224046,0.0009920587,0.0007908466,0.0004871819,0.00031736383],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011021717,0.00016798108,0.025381228,0.00067705417,0.00010484374,0.0010018237,0.0013451924,0.0005878221,0.89822793,0.0037627232,0.0013446448,0.066296645],"study_design_scores_gemma":[0.00015597205,0.0016697693,0.60007083,0.00014712708,0.0002844443,0.008306516,0.001395183,0.012319029,0.32767954,0.03971852,0.008109129,0.00014394837],"about_ca_topic_score_codex":0.0003153079,"about_ca_topic_score_gemma":0.00065255526,"teacher_disagreement_score":0.0026876004,"about_ca_system_score_codex":0.0001961101,"about_ca_system_score_gemma":0.00021627828,"threshold_uncertainty_score":0.008990943},"labels":[],"label_agreement":null},{"id":"W4399118536","doi":"10.1200/jco.2024.42.16_suppl.e13591","title":"Natural language processing for automated breast cancer recurrence detection and classification in computed tomography reports.","year":2024,"lang":"en","type":"article","venue":"Journal of Clinical Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"Provincial Health Services Authority","keywords":"Medicine; Computed tomography; Breast cancer; Radiology; Cancer; Oncology; Medical physics; Internal medicine","score_opus":0.05514140096302278,"score_gpt":0.45301280522436466,"score_spread":0.3978714042613419,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399118536","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10975127,0.0077577494,0.53273386,0.0065555377,0.0007688779,0.0049856827,0.23344389,0.09357552,0.010427685],"genre_scores_gemma":[0.13647203,0.0014915994,0.6421812,0.00070225314,0.00024610324,0.0021092251,0.21379478,0.0007531074,0.0022496632],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9953101,0.0014796064,0.0009608972,0.0012215703,0.0008678032,0.00015993866],"domain_scores_gemma":[0.9774649,0.014060668,0.0030339933,0.001462449,0.0036233899,0.00035457534],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006478284,0.0011959854,0.00075110805,0.008839999,0.0009285746,0.0023675573,0.0014030027,0.001223672,0.0065239174],"category_scores_gemma":[0.026137138,0.0004809941,0.0015702515,0.0039532077,0.0007098805,0.0029577396,0.0020443606,0.0012801363,0.0074026114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000975612,0.0005124593,0.04817314,0.007183528,0.00038719847,0.0027428283,0.0025616908,0.010072482,0.039100096,0.0051805647,0.16898833,0.71412206],"study_design_scores_gemma":[0.0002879595,0.0004905798,0.095107466,0.00177444,0.00044744712,0.0043660253,0.004177324,0.5595862,0.04735413,0.029024933,0.25713435,0.00024919247],"about_ca_topic_score_codex":0.009152487,"about_ca_topic_score_gemma":0.012361601,"teacher_disagreement_score":0.009152487,"about_ca_system_score_codex":0.0019037712,"about_ca_system_score_gemma":0.0031315978,"threshold_uncertainty_score":0.03426081},"labels":[],"label_agreement":null},{"id":"W4399368472","doi":"10.2196/62757","title":"Correction: A Multilabel Text Classifier of Cancer Literature at the Publication Level: Methods Study of Medical Text Classification","year":2024,"lang":"en","type":"erratum","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Fundamental Research Funds for the Central Universities","keywords":"Computer science; Natural language processing; Classifier (UML); Information retrieval; Artificial intelligence; Text mining","score_opus":0.05611907484497594,"score_gpt":0.4149363210982544,"score_spread":0.3588172462532785,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399368472","genre_codex":"editorial","genre_gemma":"editorial","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"editorial","genre_consensus":"editorial","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00038142072,0.0008677764,0.002767636,0.103401855,0.88307977,0.000057784127,0.0057977107,0.0012162209,0.0024298625],"genre_scores_gemma":[0.035091978,0.009385294,0.02768224,0.23029925,0.41084534,0.0007477109,0.012758078,0.008727341,0.26446277],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98535603,0.00258131,0.0028340518,0.0017403706,0.006737006,0.0007512699],"domain_scores_gemma":[0.85074633,0.053936914,0.006177805,0.0076582013,0.0778274,0.0036533948],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.011186056,0.0031277924,0.0021297785,0.008209934,0.0045975368,0.0052316203,0.004363623,0.0071520843,0.06340222],"category_scores_gemma":[0.2031063,0.0017696688,0.002169954,0.00575023,0.0036422526,0.0030667682,0.0031795444,0.010373898,0.03106294],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004214698,0.0000066647176,0.00012720516,0.00020968885,0.000029828512,0.0003545999,0.00006909544,0.0000486805,0.0000636809,0.0006134797,0.99408674,0.004348185],"study_design_scores_gemma":[0.000103365885,0.000030130372,0.0013908433,0.00096474536,0.00017278566,0.0018353661,0.00021033356,0.0007057703,0.00075312715,0.003825,0.9899053,0.000103312865],"about_ca_topic_score_codex":0.03654864,"about_ca_topic_score_gemma":0.04535124,"teacher_disagreement_score":0.98881394,"about_ca_system_score_codex":0.0048008114,"about_ca_system_score_gemma":0.010631369,"threshold_uncertainty_score":0.21210158},"labels":[],"label_agreement":null},{"id":"W4399371019","doi":"10.3138/jsp-2023-0079","title":"A Rapid Investigation of Artificial Intelligence Generated Content Footprints in Scholarly Publications","year":2024,"lang":"en","type":"article","venue":"Journal of Scholarly Publishing","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":16,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Content (measure theory); Computer science; Information retrieval; Data science; Artificial intelligence; Mathematics","score_opus":0.12928573707092703,"score_gpt":0.3083854016389009,"score_spread":0.17909966456797385,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399371019","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9570782,0.0027253632,0.014344115,0.003351789,0.00026000672,0.0004532445,0.0031818207,0.0005064715,0.018098941],"genre_scores_gemma":[0.9790471,0.00093116355,0.015161846,0.00033089286,0.00021560294,0.00027076472,0.0022936098,0.00014130796,0.0016076571],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9746055,0.0065821763,0.0035506093,0.002564059,0.011850872,0.0008466577],"domain_scores_gemma":[0.59731585,0.25403744,0.075974666,0.02185434,0.047340535,0.0034771874],"candidate_categories":["metaresearch","research_integrity"],"consensus_categories":[],"category_scores_codex":[0.021440351,0.0003196988,0.0005110355,0.024947995,0.0019071292,0.008181472,0.0009391521,0.00088245416,0.0019917304],"category_scores_gemma":[0.13846885,0.00038383092,0.00050061685,0.02558734,0.0027059899,0.010042879,0.0048724944,0.0012910229,0.0006490131],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040974005,0.00022190266,0.6074512,0.0024333824,0.00022821284,0.0020390206,0.11642786,0.0006043484,0.009650119,0.023386898,0.008249463,0.22889788],"study_design_scores_gemma":[0.000042838583,0.000276082,0.76917547,0.001767513,0.00021318172,0.0038828645,0.09503992,0.0054227747,0.008576312,0.0202312,0.09514975,0.00022200035],"about_ca_topic_score_codex":0.0012133376,"about_ca_topic_score_gemma":0.0017315783,"teacher_disagreement_score":0.99911755,"about_ca_system_score_codex":0.0019965037,"about_ca_system_score_gemma":0.0023171548,"threshold_uncertainty_score":0.11338878},"labels":[],"label_agreement":null},{"id":"W4399538170","doi":"10.1099/mgen.0.001260","title":"PHA4GE quality control contextual data tags: standardized annotations for sharing public health sequence datasets with known quality issues to facilitate testing and training","year":2024,"lang":"en","type":"article","venue":"Microbial Genomics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University; Public Health Agency of Canada; Western University; Dalhousie University; Simon Fraser University","funders":"European and Developing Countries Clinical Trials Partnership; Biotechnology and Biological Sciences Research Council; U.S. National Library of Medicine; Foreign, Commonwealth and Development Office; Bill and Melinda Gates Foundation; European Commission; Medical Research Council; Directorate for Biological Sciences; Public Health Agency; National Institutes of Health; Public Health Agency of Canada","keywords":"Quality (philosophy); Computer science; Control (management); Sequence (biology); Data quality; Information retrieval; Data science; Knowledge management; Artificial intelligence; Engineering; Biology; Operations management","score_opus":0.33853974373146933,"score_gpt":0.4206014580559228,"score_spread":0.08206171432445347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399538170","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023732305,0.0015525013,0.6540385,0.005234024,0.0029877482,0.0053029424,0.19975089,0.08362654,0.023774585],"genre_scores_gemma":[0.047734354,0.001334837,0.61866015,0.0037651851,0.0005184207,0.008209357,0.29438084,0.01707167,0.008325219],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96420306,0.011302476,0.009719279,0.0047165067,0.007950606,0.0021080533],"domain_scores_gemma":[0.8459331,0.04888119,0.018574417,0.056463793,0.0259402,0.004207415],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.065445945,0.0027646113,0.002263831,0.012082308,0.0035229735,0.007691872,0.004062884,0.0035280574,0.010098468],"category_scores_gemma":[0.12701876,0.0020088675,0.00260811,0.01059272,0.0036601597,0.00981439,0.012384986,0.00701912,0.0124715455],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.007904716,0.0012520337,0.047051832,0.010086365,0.00055940344,0.0018116687,0.011040049,0.008406655,0.08270133,0.116571,0.4274693,0.2851457],"study_design_scores_gemma":[0.0002921635,0.00029647132,0.012770612,0.0039024218,0.00022064826,0.00057717273,0.001731808,0.0064570936,0.06819549,0.035971988,0.86900645,0.0005777455],"about_ca_topic_score_codex":0.008135876,"about_ca_topic_score_gemma":0.007988164,"teacher_disagreement_score":0.9959371,"about_ca_system_score_codex":0.00499786,"about_ca_system_score_gemma":0.016079621,"threshold_uncertainty_score":0.34611535},"labels":[],"label_agreement":null},{"id":"W4399616485","doi":"10.2139/ssrn.4863974","title":"Robustly Measuring Multiple Long-Term Health Conditions Using Disparate Linked Datasets in UK Biobank","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Population and Public Health","funders":"Medical Research Council","keywords":"Biobank; Term (time); Data science; Computer science; Geography; Biology; Bioinformatics; Physics","score_opus":0.046539864908037715,"score_gpt":0.3385454665720985,"score_spread":0.29200560166406075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399616485","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7709218,0.0047823424,0.034846857,0.004056502,0.00035055698,0.00029666472,0.1784205,0.0009213522,0.0054033953],"genre_scores_gemma":[0.79232216,0.0011951235,0.0246465,0.00058844587,0.00016171693,0.0003195225,0.17918088,0.00008774053,0.0014979205],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9935126,0.001866961,0.0009172155,0.0019784004,0.0012841823,0.00044076567],"domain_scores_gemma":[0.97672075,0.013239369,0.0038150602,0.0034518929,0.0022829885,0.0004899267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005845218,0.0008738389,0.0012011799,0.006667946,0.0007141788,0.004324181,0.0011099159,0.0021903545,0.0018861962],"category_scores_gemma":[0.03527711,0.00048562937,0.0011479432,0.009356169,0.00060943706,0.0027873286,0.004047764,0.001202509,0.0014329657],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020768421,0.00045292414,0.80677474,0.0016165094,0.0028742796,0.0011423286,0.0016498517,0.018661058,0.01072,0.003745769,0.025375139,0.124910586],"study_design_scores_gemma":[0.00016762575,0.00021274075,0.90893006,0.00072887924,0.0012375512,0.0009695816,0.002311567,0.039065346,0.0054243742,0.012735257,0.0280721,0.00014489987],"about_ca_topic_score_codex":0.0248946,"about_ca_topic_score_gemma":0.032629833,"teacher_disagreement_score":0.0248946,"about_ca_system_score_codex":0.0016569857,"about_ca_system_score_gemma":0.0019583204,"threshold_uncertainty_score":0.049499393},"labels":[],"label_agreement":null},{"id":"W4399617010","doi":"10.1093/jamia/ocae143","title":"Promoting interoperability between SNOMED CT and ICD-11: lessons learned from the pilot project mapping between SNOMED CT and the ICD-11 Foundation","year":2024,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":13,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Health Information","funders":"Institut canadien d'information sur la santé; U.S. National Library of Medicine; National Institutes of Health; Bundesinstitut für Arzneimittel und Medizinprodukte; World Health Organization","keywords":"SNOMED CT; Interoperability; Systematized Nomenclature of Medicine; Foundation (evidence); Medicine; Computer science; Medical physics; World Wide Web; Terminology; Linguistics; Geography","score_opus":0.04950987139431113,"score_gpt":0.33273088483822744,"score_spread":0.2832210134439163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399617010","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.33281863,0.003597655,0.49293244,0.12792741,0.0017561602,0.009105148,0.002047224,0.0026577169,0.027157566],"genre_scores_gemma":[0.22201794,0.0015029409,0.76415235,0.005024267,0.00024632108,0.0021817815,0.0023059857,0.00058354647,0.0019848398],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9034141,0.074051015,0.0044975784,0.00442997,0.010816492,0.0027908925],"domain_scores_gemma":[0.7495518,0.14464357,0.0069565377,0.029165443,0.060465682,0.009216971],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1936233,0.0013775234,0.00090566574,0.0045276796,0.0036833424,0.0077638086,0.005974958,0.0034164006,0.0032826024],"category_scores_gemma":[0.17148612,0.00094309525,0.0014316215,0.004089464,0.005097743,0.016786158,0.015778836,0.0067516244,0.0008805123],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068926794,0.0034929002,0.03544912,0.0028253207,0.00020001936,0.0020765793,0.08376482,0.014614602,0.010855167,0.033590678,0.043119773,0.7693218],"study_design_scores_gemma":[0.0010325772,0.0055551543,0.050329998,0.013678009,0.00040727464,0.0044882405,0.20710582,0.06932542,0.031100359,0.16068967,0.45532787,0.0009596661],"about_ca_topic_score_codex":0.017616846,"about_ca_topic_score_gemma":0.02134611,"teacher_disagreement_score":0.1936233,"about_ca_system_score_codex":0.0056744693,"about_ca_system_score_gemma":0.031602513,"threshold_uncertainty_score":0.99440604},"labels":[],"label_agreement":null},{"id":"W4399668246","doi":"10.1145/3661167.3661172","title":"The Promise and Challenges of Using LLMs to Accelerate the Screening Process of Systematic Reviews","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":22,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University","funders":"Strategic Research Council; Killam Trusts","keywords":"Process (computing); Systematic review; Computer science; Risk analysis (engineering); Medicine; MEDLINE; Political science","score_opus":0.15905957729080952,"score_gpt":0.36787879599789824,"score_spread":0.20881921870708872,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399668246","genre_codex":"methods","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":"methods","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12713918,0.073887445,0.5941712,0.14006916,0.0035082903,0.018091021,0.0040674806,0.026980633,0.012085577],"genre_scores_gemma":[0.07990169,0.0053403648,0.904202,0.0030926412,0.00055650267,0.0052639623,0.0006763358,0.0005581839,0.00040830846],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.4662976,0.4791699,0.026056275,0.009178506,0.017963301,0.0013344152],"domain_scores_gemma":[0.030512936,0.89279556,0.023387115,0.031883627,0.019547725,0.0018730079],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.47477266,0.0033271383,0.0047216597,0.011823635,0.0024069564,0.012730881,0.004738777,0.0042898757,0.0076915044],"category_scores_gemma":[0.8019078,0.0038312515,0.0065879934,0.010673674,0.0034703873,0.018734943,0.009267641,0.0064626723,0.0035213348],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0063082604,0.0008824566,0.012641102,0.076018706,0.004505684,0.00026014788,0.015294042,0.01162647,0.010588576,0.009831676,0.02494404,0.82709897],"study_design_scores_gemma":[0.014316255,0.015645675,0.0876583,0.1077236,0.015565727,0.0029302693,0.013051776,0.23609221,0.04296932,0.20408253,0.25618297,0.0037812702],"about_ca_topic_score_codex":0.005860028,"about_ca_topic_score_gemma":0.015884738,"teacher_disagreement_score":0.5252273,"about_ca_system_score_codex":0.009207926,"about_ca_system_score_gemma":0.030901976,"threshold_uncertainty_score":0.6476988},"labels":[],"label_agreement":null},{"id":"W4399936079","doi":"10.1145/3673649","title":"Will AI Flood Us with Irrelevant Papers?","year":2024,"lang":"en","type":"article","venue":"Communications of the ACM","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Montréal","funders":"","keywords":"Computer science; Raising (metalworking); Citation; Data science; Flood myth; World Wide Web; History; Engineering; Archaeology","score_opus":0.01744835621752669,"score_gpt":0.2929129966476965,"score_spread":0.2754646404301698,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4399936079","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0146582695,0.031699482,0.040687185,0.8699178,0.009323816,0.00011363255,0.0010875283,0.0010856164,0.03142669],"genre_scores_gemma":[0.6291464,0.036736473,0.07029853,0.17404789,0.05087425,0.00047325384,0.0019062368,0.0019068912,0.034610096],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.92177033,0.047367577,0.0044729365,0.006078258,0.017767899,0.0025430715],"domain_scores_gemma":[0.31975663,0.51151055,0.034586776,0.035946988,0.08279833,0.015400769],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13658915,0.0008922742,0.0028733395,0.011878655,0.005953998,0.021889934,0.003975928,0.008805768,0.017176451],"category_scores_gemma":[0.54772353,0.00090619567,0.0012489918,0.014002845,0.01131298,0.04491411,0.004687512,0.007946162,0.010064666],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00062903075,0.00022409968,0.018554697,0.0014208349,0.00063518976,0.00036066675,0.0028412172,0.0015993483,0.00061630737,0.23147689,0.33691332,0.40472835],"study_design_scores_gemma":[0.00017438139,0.000118884316,0.0059941383,0.000795098,0.0001957146,0.0007016721,0.0020273314,0.0055134655,0.00079456426,0.8131297,0.1703911,0.00016387881],"about_ca_topic_score_codex":0.0053837197,"about_ca_topic_score_gemma":0.005032304,"teacher_disagreement_score":0.86341083,"about_ca_system_score_codex":0.0055089816,"about_ca_system_score_gemma":0.009635263,"threshold_uncertainty_score":0.7223611},"labels":[],"label_agreement":null},{"id":"W4400058808","doi":"10.1101/2024.06.21.599225","title":"Using fingerprinting as a testbed for strategies to improve reproducibility of functional connectivity","year":2024,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Reproducibility; Testbed; Computer science; Functional connectivity; Psychology; Mathematics; Statistics; Computer network; Neuroscience","score_opus":0.03485550796422546,"score_gpt":0.2895755774558948,"score_spread":0.25472006949166937,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400058808","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07973565,0.0004671875,0.9114253,0.0010672832,0.000184988,0.00035552977,0.0012023201,0.0036729577,0.0018887682],"genre_scores_gemma":[0.42688182,0.00024980554,0.5672745,0.0004814609,0.0001347461,0.0012307087,0.002163609,0.0010208719,0.0005624517],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9673557,0.020391244,0.001965334,0.006283415,0.0032683997,0.00073583575],"domain_scores_gemma":[0.86062425,0.06577443,0.011556295,0.05129771,0.00939598,0.0013512602],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.06684251,0.0017783518,0.0015037275,0.0035356293,0.0017280278,0.0052008643,0.00253632,0.0024903836,0.0019973342],"category_scores_gemma":[0.1830131,0.0010371266,0.0017782858,0.0034262743,0.0037124148,0.004374384,0.004956797,0.0032276777,0.001133347],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020222042,0.0010303891,0.1811385,0.0021214285,0.0037893825,0.00077254314,0.0052490137,0.10263236,0.20123155,0.057979085,0.011958583,0.4300751],"study_design_scores_gemma":[0.00042809866,0.0021976817,0.13318925,0.00060238165,0.0012372398,0.0016110817,0.0012591276,0.37926543,0.20933719,0.23793933,0.032235596,0.0006975994],"about_ca_topic_score_codex":0.0018380372,"about_ca_topic_score_gemma":0.0017859185,"teacher_disagreement_score":0.9331575,"about_ca_system_score_codex":0.0012276821,"about_ca_system_score_gemma":0.0023048366,"threshold_uncertainty_score":0.3535012},"labels":[],"label_agreement":null},{"id":"W4400140828","doi":"10.1093/bioinformatics/btae237","title":"Predicting protein functions using positive-unlabeled ranking with ontology-based priors","year":2024,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"King Abdullah University of Science and Technology","keywords":"Computer science; Prior probability; Benchmark (surveying); Classifier (UML); Ranking (information retrieval); Artificial intelligence; Machine learning; Data mining; Source code; Gene ontology; Function (biology); Ontology; Pattern recognition (psychology); Bayesian probability","score_opus":0.015018811350751586,"score_gpt":0.2557952577102794,"score_spread":0.2407764463595278,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400140828","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15360378,0.0021703406,0.80591136,0.0013277704,0.00018800556,0.00031375364,0.011715661,0.016670913,0.008098375],"genre_scores_gemma":[0.58115494,0.00053224416,0.37714922,0.0006529605,0.0002530995,0.00032585798,0.034498088,0.000983884,0.004449745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9969439,0.00084357854,0.0001438835,0.0009123813,0.0009419553,0.00021426743],"domain_scores_gemma":[0.99402195,0.002835857,0.00049499265,0.0014556252,0.0009391596,0.0002523677],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033715232,0.0017754812,0.001682004,0.004324429,0.0010423572,0.0020429757,0.0029522325,0.002174607,0.002093424],"category_scores_gemma":[0.013091617,0.0004940143,0.0013954185,0.0024608255,0.0009447964,0.0035875363,0.0019864966,0.0018802519,0.0023965703],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013934359,0.0017441652,0.035308644,0.0008947916,0.0003836644,0.0004503644,0.00029970478,0.32199413,0.027199559,0.026412996,0.06615528,0.5177633],"study_design_scores_gemma":[0.000039123173,0.0000760882,0.0023496267,0.000031500575,0.000034878,0.00013963114,0.000051282408,0.95115995,0.0055457223,0.036375854,0.0041674674,0.000028820754],"about_ca_topic_score_codex":0.005282251,"about_ca_topic_score_gemma":0.013880524,"teacher_disagreement_score":0.005282251,"about_ca_system_score_codex":0.0015582555,"about_ca_system_score_gemma":0.0017158174,"threshold_uncertainty_score":0.01783055},"labels":[],"label_agreement":null},{"id":"W4400166362","doi":"10.1016/b978-0-323-95502-7.00062-2","title":"Biological and Medical Ontologies: PRotein Ontology (PRO)","year":2024,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Ontology; Computer science; Open Biomedical Ontologies; Computational biology; Information retrieval; Data science; World Wide Web; Process ontology; Biology; Ontology alignment; Semantic Web; Epistemology; Philosophy","score_opus":0.024276602486349705,"score_gpt":0.27376993659290105,"score_spread":0.24949333410655133,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400166362","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016689571,0.032501433,0.41486567,0.011542092,0.0067252265,0.00026086843,0.025644258,0.014813373,0.4919782],"genre_scores_gemma":[0.009722308,0.03828034,0.287621,0.0067471615,0.0038270683,0.0005214123,0.050561093,0.00789193,0.5948277],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9995734,0.000053307667,0.000050464623,0.00007098764,0.00022845724,0.000023363096],"domain_scores_gemma":[0.9992612,0.00033442,0.000059505393,0.000106177606,0.00016439214,0.00007420971],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00095603475,0.0016420003,0.0009076694,0.0056649046,0.00059994886,0.0050979075,0.00095362985,0.0012631334,0.10073772],"category_scores_gemma":[0.0020438072,0.0008602106,0.0007500337,0.009719476,0.0009143762,0.009436655,0.0024931189,0.0021931066,0.0879776],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017833387,0.000027362797,0.00006915836,0.0006496229,0.0000140666225,0.00009969237,0.00022826057,0.000374838,0.0021311673,0.082517855,0.51957923,0.39429086],"study_design_scores_gemma":[0.000003986998,0.000003848619,0.00010485399,0.00014853422,0.0000061047535,0.0002645622,0.000046852747,0.0005697753,0.0002885082,0.028541775,0.970014,0.000007302505],"about_ca_topic_score_codex":0.002162821,"about_ca_topic_score_gemma":0.003126282,"teacher_disagreement_score":0.10073772,"about_ca_system_score_codex":0.001154655,"about_ca_system_score_gemma":0.0009873168,"threshold_uncertainty_score":0.33700126},"labels":[],"label_agreement":null},{"id":"W4400606045","doi":"10.1038/s41598-024-65645-6","title":"Identifying symptom etiologies using syntactic patterns and large language models","year":2024,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University Health Centre","funders":"H2020 European Research Council; Israel Science Foundation; European Commission","keywords":"Computer science; Etiology; Bootstrapping (finance); Natural language processing; Medical diagnosis; Artificial intelligence; Generative grammar; Reliability (semiconductor); Pipeline (software); Machine learning; Pathology; Medicine; Programming language","score_opus":0.03139023252115055,"score_gpt":0.3188436165696493,"score_spread":0.2874533840484988,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400606045","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043495562,0.0016307702,0.93371755,0.0028874392,0.00008990303,0.00046767507,0.008786474,0.0052923243,0.0036323208],"genre_scores_gemma":[0.3450043,0.001728437,0.6306818,0.0006029609,0.00017656725,0.00053061615,0.019409934,0.00054154394,0.0013238037],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9967456,0.0010883374,0.00044623553,0.0007116815,0.0008940952,0.000113997565],"domain_scores_gemma":[0.98501647,0.011354629,0.0016457979,0.0010075567,0.00080291973,0.00017266469],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038016739,0.0013991046,0.00066225324,0.00886129,0.0010033948,0.0031751674,0.0015519498,0.0011095247,0.0024298213],"category_scores_gemma":[0.015852211,0.00077649753,0.0032311743,0.0044165286,0.001182299,0.0048783612,0.0030463473,0.0023302832,0.001768101],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00049113476,0.000552193,0.081931286,0.0031298648,0.0012480912,0.006732504,0.0057527483,0.05613432,0.023576159,0.11388006,0.023479227,0.68309236],"study_design_scores_gemma":[0.00009186152,0.00015482403,0.017592829,0.0006503756,0.0008654342,0.003395038,0.0018623527,0.58431137,0.012024672,0.32658046,0.052312993,0.00015776644],"about_ca_topic_score_codex":0.0039794017,"about_ca_topic_score_gemma":0.008396247,"teacher_disagreement_score":0.00886129,"about_ca_system_score_codex":0.0011646359,"about_ca_system_score_gemma":0.0025764387,"threshold_uncertainty_score":0.020105422},"labels":[],"label_agreement":null},{"id":"W4400866246","doi":"10.1101/2024.07.16.603812","title":"SciMind: A Multimodal Mixture-of-Experts Model for Advancing Pharmaceutical Sciences","year":2024,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"PROTO Manufacturing (Canada)","funders":"","keywords":"Pharmaceutical sciences; Computer science; Data science; Medicine; Pharmacology","score_opus":0.0286781551357841,"score_gpt":0.30867759767643504,"score_spread":0.27999944254065096,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400866246","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020902855,0.0009307765,0.9704763,0.0013163234,0.000121146164,0.00013456262,0.00092891697,0.002687084,0.0025020484],"genre_scores_gemma":[0.5402198,0.0006560475,0.44658118,0.0011747031,0.00022792965,0.00061455846,0.002536944,0.00034034887,0.007648519],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99908566,0.00039413242,0.00004520968,0.00022670753,0.00016581059,0.00008243347],"domain_scores_gemma":[0.9982173,0.0012114426,0.00009960109,0.000091489106,0.0002690348,0.000111094196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027206657,0.0012144771,0.0010205309,0.0012013866,0.00049850927,0.0011413085,0.002546047,0.002069752,0.0035697306],"category_scores_gemma":[0.0049934727,0.00076436443,0.0015355692,0.00086001604,0.00072873454,0.0015341421,0.0020276466,0.0027216065,0.0010092888],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028458735,0.0001072322,0.0009102262,0.00010633999,0.000113397145,0.00008424766,0.00006542981,0.929901,0.0011730684,0.007594053,0.0052511697,0.054409277],"study_design_scores_gemma":[0.000008487568,0.000009009904,0.00002857196,0.0000035806775,0.0000044845074,0.0000048340876,0.0000019565073,0.99710935,0.00018199562,0.002269304,0.0003754603,0.0000029984803],"about_ca_topic_score_codex":0.01129452,"about_ca_topic_score_gemma":0.012717158,"teacher_disagreement_score":0.01129452,"about_ca_system_score_codex":0.0018490755,"about_ca_system_score_gemma":0.0019685333,"threshold_uncertainty_score":0.02245754},"labels":[],"label_agreement":null},{"id":"W4400895718","doi":"10.2196/57853","title":"PCEtoFHIR: Decomposition of Postcoordinated SNOMED CT Expressions for Storage as HL7 FHIR Resources","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"SNOMED CT; Computer science; Information retrieval; Database; Data mining; Terminology","score_opus":0.011984129367489162,"score_gpt":0.3342222578306032,"score_spread":0.32223812846311406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400895718","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008943451,0.00019773463,0.9172971,0.00063699565,0.00022114754,0.0009598619,0.007912675,0.043824658,0.020006314],"genre_scores_gemma":[0.06847128,0.00033897228,0.8750361,0.0006030549,0.0001311877,0.0007445825,0.029275367,0.008715622,0.016683795],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99722195,0.00043326756,0.00048684605,0.0005348842,0.001054177,0.00026887064],"domain_scores_gemma":[0.9960361,0.0010369197,0.00033757376,0.0016201979,0.00083179504,0.0001374717],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004334919,0.0012451689,0.00075667904,0.0032623082,0.0010404711,0.0049702474,0.0022825925,0.0009533573,0.018624786],"category_scores_gemma":[0.008932778,0.00080916274,0.002206025,0.0023872468,0.0013697161,0.0057149366,0.0044212486,0.0018389269,0.008181434],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013651911,0.000417527,0.0042344737,0.0011634291,0.00016453651,0.002070053,0.0047638947,0.00878589,0.045219,0.26867306,0.1405508,0.5225922],"study_design_scores_gemma":[0.00021484286,0.00020368306,0.0041006007,0.0004627556,0.0001595026,0.001823447,0.001574617,0.09880295,0.12998576,0.10111631,0.6612559,0.00029957548],"about_ca_topic_score_codex":0.008133906,"about_ca_topic_score_gemma":0.0071830545,"teacher_disagreement_score":0.018624786,"about_ca_system_score_codex":0.0022012785,"about_ca_system_score_gemma":0.003983303,"threshold_uncertainty_score":0.062306106},"labels":[],"label_agreement":null},{"id":"W4400918589","doi":"10.2196/54653","title":"Accelerating Evidence Synthesis in Observational Studies: Development of a Living Natural Language Processing–Assisted Intelligent Systematic Literature Review System","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Systematic review; Observational study; Machine learning; Data extraction; Artificial intelligence; Natural language processing; MEDLINE; Medicine","score_opus":0.08055717155902066,"score_gpt":0.38488936394622786,"score_spread":0.3043321923872072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400918589","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0070540872,0.0019159726,0.85803926,0.0049540484,0.0005055684,0.012791166,0.009961719,0.10192519,0.0028529835],"genre_scores_gemma":[0.006964002,0.0002815372,0.9850595,0.00044698475,0.000047437155,0.0046533193,0.0017630564,0.0004645325,0.00031955453],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9367003,0.03365148,0.01670132,0.0067084944,0.005766843,0.0004716366],"domain_scores_gemma":[0.74504095,0.18418172,0.019582486,0.02331955,0.025161669,0.002713682],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13434334,0.002269068,0.0033065823,0.017916137,0.0020007847,0.0075706216,0.0035522154,0.002383341,0.009539582],"category_scores_gemma":[0.22166681,0.0023878645,0.004893487,0.009388008,0.0011829479,0.006742093,0.0085865855,0.0028406018,0.005387743],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014739339,0.00045350078,0.008582455,0.037318006,0.0043519945,0.001339248,0.008885429,0.011270509,0.014022664,0.019120337,0.09517337,0.7980085],"study_design_scores_gemma":[0.0037181205,0.0012936897,0.013317597,0.020294635,0.0070234607,0.0019303517,0.0037888186,0.38000062,0.03935488,0.09611973,0.43185577,0.0013023512],"about_ca_topic_score_codex":0.0026784842,"about_ca_topic_score_gemma":0.0059658284,"teacher_disagreement_score":0.8656567,"about_ca_system_score_codex":0.0035736437,"about_ca_system_score_gemma":0.017236806,"threshold_uncertainty_score":0.710484},"labels":[],"label_agreement":null},{"id":"W4400935684","doi":"10.3233/shti240162","title":"Going Beyond Surface Language: An Exploratory Evaluation of Nursing Ontology Mappings","year":2024,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Terminology; Computer science; Principle of compositionality; Ontology; Systematized Nomenclature of Medicine; Viewpoints; SNOMED CT; Natural language processing; Automation; Data science; Artificial intelligence; Information retrieval; Linguistics; Epistemology","score_opus":0.08094757640278819,"score_gpt":0.43491538658740514,"score_spread":0.353967810184617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4400935684","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93109375,0.00040222687,0.054252725,0.0006499362,0.000064475345,0.00418392,0.00065740524,0.00031145298,0.008384068],"genre_scores_gemma":[0.8262445,0.0002865158,0.16618463,0.0003065236,0.000024260144,0.0048729316,0.00081717415,0.00022479492,0.0010387534],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9073692,0.07486037,0.0065442408,0.0023602264,0.007687907,0.0011781055],"domain_scores_gemma":[0.6589893,0.29422948,0.007952396,0.011533349,0.026145931,0.0011495348],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.081276275,0.0009302487,0.00085792557,0.00428449,0.002230494,0.003985116,0.0016220377,0.0013100947,0.0020133934],"category_scores_gemma":[0.26441088,0.00047659723,0.0011315898,0.004276706,0.0025406,0.0056692925,0.00546811,0.0012878967,0.00040021027],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0097142095,0.004771666,0.07388904,0.0089834295,0.000719782,0.0034736996,0.3074154,0.017080644,0.040842853,0.02848443,0.00546139,0.49916348],"study_design_scores_gemma":[0.0031683003,0.032665983,0.13146557,0.00647062,0.0020112284,0.0064197686,0.38160446,0.1388359,0.09726931,0.07862747,0.12036218,0.0010992193],"about_ca_topic_score_codex":0.0027179709,"about_ca_topic_score_gemma":0.0039200783,"teacher_disagreement_score":0.081276275,"about_ca_system_score_codex":0.0023916992,"about_ca_system_score_gemma":0.002563331,"threshold_uncertainty_score":0.42983514},"labels":[],"label_agreement":null},{"id":"W4401034543","doi":"10.1177/14604582241267792","title":"A novel technology for harmonizing and analyzing cancer data. Observations from integrating health connect in Newfoundland and Labrador, Canada","year":2024,"lang":"en","type":"article","venue":"Health Informatics Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Memorial University of Newfoundland; Newfoundland and Labrador Centre for Applied Health Research","funders":"","keywords":"Cancer; Health data; Data science; Computer science; Geography; Knowledge management; Health care; Medicine; Political science; Internal medicine","score_opus":0.10480036553276482,"score_gpt":0.3553497999984447,"score_spread":0.25054943446567984,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401034543","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4173348,0.0049435315,0.23910896,0.04349502,0.0005195853,0.0027131173,0.13623665,0.017909298,0.1377391],"genre_scores_gemma":[0.6139934,0.0024594853,0.2861711,0.003482673,0.00013434977,0.00073157536,0.0677828,0.0008529922,0.024391565],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99451905,0.0009118805,0.00040718185,0.0009735853,0.0026319802,0.00055632234],"domain_scores_gemma":[0.990639,0.0019805976,0.0008655049,0.0015480008,0.0043875254,0.0005794011],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005539975,0.000590611,0.00034305218,0.0038205436,0.002658503,0.002900255,0.0011475486,0.00044234117,0.0025997062],"category_scores_gemma":[0.01272632,0.00043538044,0.00069346005,0.011014277,0.0012926488,0.0014349656,0.0026811694,0.00068180374,0.0005420209],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00088529807,0.00019090451,0.3149122,0.00080930284,0.00048631092,0.00068962603,0.006446752,0.016422503,0.012279627,0.025608016,0.15611319,0.4651562],"study_design_scores_gemma":[0.00017518846,0.0001968054,0.39387253,0.00040823463,0.00038201382,0.00060693506,0.008027691,0.040864285,0.015965216,0.008544345,0.53074336,0.0002134148],"about_ca_topic_score_codex":0.96591437,"about_ca_topic_score_gemma":0.9752695,"teacher_disagreement_score":0.03408563,"about_ca_system_score_codex":0.028664524,"about_ca_system_score_gemma":0.044409595,"threshold_uncertainty_score":0.20797664},"labels":[],"label_agreement":null},{"id":"W4401042030","doi":"10.18653/v1/2024.semeval-1.239","title":"CLaC at SemEval-2024 Task 2: Faithful Clinical Trial Inference","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"SemEval; Inference; Computer science; Task (project management); Natural language processing; Artificial intelligence","score_opus":0.06425536072042741,"score_gpt":0.40959676904826114,"score_spread":0.3453414083278337,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042030","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027658073,0.0029893892,0.7611367,0.0077644517,0.0016370591,0.006591864,0.095363416,0.07356953,0.02328954],"genre_scores_gemma":[0.13908581,0.00050111505,0.6997485,0.0031474656,0.00054548064,0.004348443,0.13813283,0.006018802,0.008471609],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9742276,0.015518481,0.0017307586,0.004483924,0.0033504213,0.00068882],"domain_scores_gemma":[0.89794314,0.07488727,0.002488769,0.0138777755,0.0089500975,0.0018529447],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03530348,0.003993926,0.0020352479,0.00577786,0.0023273483,0.005285045,0.0051196082,0.0067642676,0.03473715],"category_scores_gemma":[0.14849642,0.001412257,0.0043866434,0.0023953347,0.0018573357,0.004842472,0.0067056706,0.0061782254,0.012991819],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030912382,0.0012709964,0.0106630325,0.0069133462,0.0015372701,0.0015611458,0.0016215739,0.06951039,0.010524326,0.03823089,0.53796446,0.31711137],"study_design_scores_gemma":[0.0022053844,0.0009793267,0.005044842,0.001484604,0.00072251284,0.0019459566,0.0007164388,0.48483407,0.031800695,0.13738485,0.33259138,0.0002899208],"about_ca_topic_score_codex":0.008939264,"about_ca_topic_score_gemma":0.020769324,"teacher_disagreement_score":0.03530348,"about_ca_system_score_codex":0.003126526,"about_ca_system_score_gemma":0.009626063,"threshold_uncertainty_score":0.18670493},"labels":[],"label_agreement":null},{"id":"W4401042733","doi":"10.18653/v1/2024.semeval-1.26","title":"iML at SemEval-2024 Task 2: Safe Biomedical Natural Language Interference for Clinical Trials with LLM Based Ensemble Inferencing","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"SemEval; Computer science; Task (project management); Natural language processing; Artificial intelligence; Natural language; Interference (communication); Natural (archaeology); Engineering; Biology","score_opus":0.08176553105415232,"score_gpt":0.42913818955879035,"score_spread":0.34737265850463805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401042733","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042149812,0.003414886,0.83716214,0.013269104,0.0016167712,0.002768194,0.03987087,0.045464184,0.014283981],"genre_scores_gemma":[0.26433453,0.0006770795,0.65500665,0.0044005937,0.00064226554,0.0028080323,0.06182501,0.003302654,0.00700318],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9665063,0.023525795,0.0017085702,0.0041479575,0.0033272258,0.00078414433],"domain_scores_gemma":[0.9109158,0.07009332,0.0020287824,0.009924797,0.0055697123,0.0014676035],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.040437784,0.0026512723,0.0019880598,0.0023847667,0.0017665976,0.004007995,0.005683573,0.0055888356,0.01674519],"category_scores_gemma":[0.11027913,0.0010500391,0.0034105934,0.0013965394,0.0020482868,0.0051111192,0.008371103,0.0067035235,0.0057768333],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038012126,0.00173726,0.0069856253,0.0063838596,0.001649857,0.0025531005,0.0026040268,0.11746621,0.023724519,0.06476979,0.30408475,0.46423975],"study_design_scores_gemma":[0.001272305,0.0009633601,0.0024169285,0.00064451207,0.0006432029,0.0013299073,0.0005713479,0.6480435,0.0523715,0.1601529,0.13133359,0.00025685682],"about_ca_topic_score_codex":0.004378188,"about_ca_topic_score_gemma":0.007771085,"teacher_disagreement_score":0.040437784,"about_ca_system_score_codex":0.0035867041,"about_ca_system_score_gemma":0.008456159,"threshold_uncertainty_score":0.21385801},"labels":[],"label_agreement":null},{"id":"W4401184890","doi":"10.1098/rspb.2024.0423","title":"The changing landscape of text mining: a review of approaches for ecology and evolution","year":2024,"lang":"en","type":"review","venue":"Proceedings of the Royal Society B Biological Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"Medical Research Council; Natural Sciences and Engineering Research Council of Canada; University of Glasgow","keywords":"Computer science; Data science; Evolutionary ecology; Process (computing); Ecology; Artificial intelligence; Fitness landscape; Biology; Sociology","score_opus":0.07081132361903111,"score_gpt":0.3167234877637411,"score_spread":0.24591216414470998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401184890","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007095335,0.9954982,0.0012529219,0.002257382,0.0002856422,0.000011341778,0.0000391324,0.000024437268,0.00055998657],"genre_scores_gemma":[0.0005631903,0.99528307,0.002404221,0.0010489075,0.00039118383,0.000024818273,0.000057539466,0.00001089748,0.0002162861],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99777836,0.00074742245,0.0004510579,0.00029207958,0.00065032317,0.00008083437],"domain_scores_gemma":[0.9808153,0.015872903,0.00094582274,0.00036103564,0.0016741806,0.0003308449],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005828605,0.001215712,0.0023641023,0.010030673,0.0006335899,0.0033900652,0.0022362638,0.0023533236,0.0031020716],"category_scores_gemma":[0.01292832,0.0006253282,0.0016071146,0.0115093365,0.0027783278,0.0078082313,0.0017016326,0.0030812097,0.0020516503],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003401011,0.000031356125,0.00021084667,0.0299419,0.00017289394,0.00007868943,0.00021486176,0.00033028648,0.00040988927,0.010031534,0.026552254,0.93199146],"study_design_scores_gemma":[0.000022195116,0.00006897626,0.0016256972,0.034300692,0.0002322277,0.00079407694,0.00032120317,0.00041548407,0.00040735502,0.020915607,0.94082147,0.000074925374],"about_ca_topic_score_codex":0.0027575071,"about_ca_topic_score_gemma":0.0040773614,"teacher_disagreement_score":0.010030673,"about_ca_system_score_codex":0.0020719022,"about_ca_system_score_gemma":0.00410815,"threshold_uncertainty_score":0.03082496},"labels":[],"label_agreement":null},{"id":"W4401625984","doi":"10.20944/preprints202408.0966.v1","title":"Enhancing the Interpretability of Malaria and Typhoid Diagnosis with Explainable AI and Large Language Models","year":2024,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mount Royal University","funders":"","keywords":"Interpretability; Typhoid fever; Malaria; Computer science; Artificial intelligence; Plasmodium falciparum; Natural language processing; Medicine; Virology; Immunology","score_opus":0.03433271375113809,"score_gpt":0.3200210188595128,"score_spread":0.2856883051083747,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401625984","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22738929,0.0015081959,0.7380417,0.009832857,0.0003321309,0.00035617358,0.0034415107,0.01314871,0.005949401],"genre_scores_gemma":[0.825575,0.00038945157,0.16761792,0.00090285647,0.00012122938,0.00020093455,0.002801818,0.00037997888,0.0020107208],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978377,0.0014210922,0.00010222295,0.0003704532,0.00016362187,0.00010485968],"domain_scores_gemma":[0.98153955,0.016429104,0.0005914956,0.0005557809,0.0006834787,0.0002006203],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0046401853,0.0012998811,0.00054549705,0.0011922244,0.00043189697,0.003072814,0.0013412123,0.001166153,0.003981761],"category_scores_gemma":[0.027988417,0.00041451858,0.001769919,0.00060743064,0.0005835751,0.0033267762,0.0018564431,0.0030920252,0.0010774287],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013203331,0.00049450155,0.034915447,0.0009830999,0.0005768878,0.0011865223,0.0044328542,0.49124214,0.009608857,0.022415414,0.019318582,0.41350535],"study_design_scores_gemma":[0.000023821689,0.000049167065,0.0013478375,0.000044391174,0.00005974898,0.00007008786,0.00016127611,0.9800811,0.0017643485,0.014594796,0.0017710166,0.000032542226],"about_ca_topic_score_codex":0.010232633,"about_ca_topic_score_gemma":0.01149263,"teacher_disagreement_score":0.010232633,"about_ca_system_score_codex":0.0015821492,"about_ca_system_score_gemma":0.001519378,"threshold_uncertainty_score":0.024539948},"labels":[],"label_agreement":null},{"id":"W4401689426","doi":"10.1186/s13326-024-00315-0","title":"Concretizing plan specifications as realizables within the OBO foundry","year":2024,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science; Plan (archaeology); Foundry; Engineering management; Knowledge management; Software engineering; Manufacturing engineering; Data science; Systems engineering; Management science; Mechanical engineering; Engineering","score_opus":0.04353713051971736,"score_gpt":0.3029309165570618,"score_spread":0.2593937860373444,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401689426","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027138663,0.00015716151,0.9496688,0.0018857486,0.000091192094,0.00033394527,0.0004132303,0.0020761879,0.01823506],"genre_scores_gemma":[0.33601698,0.0003402093,0.65496665,0.00052760233,0.00004798825,0.00039416904,0.001185656,0.0008514407,0.005669236],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99245125,0.0023735785,0.0010393158,0.0009693135,0.0025032551,0.00066322],"domain_scores_gemma":[0.98092514,0.007551302,0.00224526,0.006630463,0.0022086296,0.00043917462],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009732684,0.00073454797,0.00052161043,0.0016917911,0.0018443958,0.005469106,0.0022452255,0.0018977626,0.0038517541],"category_scores_gemma":[0.023505697,0.00083510735,0.0022825394,0.0015096306,0.007891637,0.011373664,0.0062977304,0.0032701541,0.00063859683],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054180236,0.000032533342,0.0022620908,0.00015362968,0.000030918553,0.00035071766,0.003678658,0.009343143,0.0016897669,0.95205176,0.0018278954,0.028524637],"study_design_scores_gemma":[0.000046985082,0.00006697794,0.0011998513,0.00036802335,0.00015063141,0.0005317171,0.0020671524,0.081520826,0.013570617,0.7567019,0.1436796,0.00009579655],"about_ca_topic_score_codex":0.021901438,"about_ca_topic_score_gemma":0.021355191,"teacher_disagreement_score":0.021901438,"about_ca_system_score_codex":0.0047487277,"about_ca_system_score_gemma":0.009576068,"threshold_uncertainty_score":0.05147195},"labels":[],"label_agreement":null},{"id":"W4401821632","doi":"10.1016/j.simpa.2024.100699","title":"PostgREST Data Provider for React-Admin: Bootstrap the creation of user interfaces on top of PostgreSQL databases","year":2024,"lang":"en","type":"article","venue":"Software Impacts","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Bundesministerium für Bildung, Wissenschaft, Forschung und Technologie; Bundesministerium für Bildung und Forschung","keywords":"Database; Computer science; World Wide Web","score_opus":0.08301828793632884,"score_gpt":0.3862413130828631,"score_spread":0.3032230251465342,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401821632","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007666746,0.00021492665,0.2569906,0.0010933524,0.0004283345,0.00066162687,0.009622332,0.7078303,0.015491863],"genre_scores_gemma":[0.22379294,0.00078013475,0.29080802,0.006505156,0.00057743327,0.0023585388,0.0907037,0.33102542,0.05344862],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9925236,0.0015139285,0.0007011402,0.0017409356,0.0028015226,0.000718831],"domain_scores_gemma":[0.9789179,0.0045039253,0.0008177678,0.010851962,0.0031963177,0.0017120481],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014232825,0.0023330164,0.0014475405,0.0019572116,0.001215201,0.006225359,0.0047865934,0.0016533589,0.04696525],"category_scores_gemma":[0.026618596,0.0024092807,0.0019883371,0.0017716704,0.0017335693,0.009546987,0.011365711,0.0046793055,0.04725207],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0056984005,0.00074691954,0.009745681,0.0010325877,0.00035864496,0.0015992102,0.003936881,0.0037192751,0.030264812,0.046287704,0.7314011,0.16520882],"study_design_scores_gemma":[0.00083036633,0.00029433318,0.0050521167,0.00038088308,0.00013340259,0.0007391615,0.0010365252,0.07177646,0.0749565,0.028765375,0.8155042,0.0005307211],"about_ca_topic_score_codex":0.0040857815,"about_ca_topic_score_gemma":0.0023384965,"teacher_disagreement_score":0.04696525,"about_ca_system_score_codex":0.0019318076,"about_ca_system_score_gemma":0.0032264045,"threshold_uncertainty_score":0.15711439},"labels":[],"label_agreement":null},{"id":"W4401823954","doi":"10.3233/shti240659","title":"Linking Health Terminologies: A Unified Approach to the WHO Family of International Classifications","year":2024,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Institute for Health Information","funders":"World Health Organization","keywords":"Interoperability; Computer science; Linkage (software); Informatics; Health informatics; Data science; Knowledge management; World Wide Web; Health care; Engineering","score_opus":0.08491963788981843,"score_gpt":0.38637381359814804,"score_spread":0.3014541757083296,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401823954","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020624478,0.0035721597,0.947388,0.008348458,0.0012684459,0.00051269174,0.0009732173,0.00087965967,0.034994937],"genre_scores_gemma":[0.026814505,0.0033770339,0.96052414,0.0018293097,0.00092007994,0.0009527742,0.0023494493,0.00032056787,0.0029121018],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97610915,0.0122739365,0.0038630627,0.0018692401,0.005160096,0.00072448415],"domain_scores_gemma":[0.97172767,0.011231959,0.0025720445,0.005461351,0.008067259,0.0009397065],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.034798056,0.0010900892,0.0016727048,0.032517824,0.0035721736,0.012510324,0.0037734034,0.002942375,0.0029983132],"category_scores_gemma":[0.040145785,0.0008081949,0.0022396126,0.024428345,0.008235027,0.020174801,0.0069025634,0.0049445475,0.0020039235],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011959015,0.000019986672,0.00085544,0.000256615,0.000028602351,0.000111316076,0.0026331046,0.00077076,0.00039132248,0.9154325,0.010308588,0.06917982],"study_design_scores_gemma":[0.000016928361,0.000041124127,0.0012194077,0.0013895758,0.00009181063,0.00050613633,0.0027232196,0.0062769935,0.0006732802,0.4788143,0.50817287,0.00007436931],"about_ca_topic_score_codex":0.00811826,"about_ca_topic_score_gemma":0.0057391315,"teacher_disagreement_score":0.034798056,"about_ca_system_score_codex":0.0056731994,"about_ca_system_score_gemma":0.015635045,"threshold_uncertainty_score":0.1840319},"labels":[],"label_agreement":null},{"id":"W4401824717","doi":"10.3233/shti240647","title":"Enhancing Drug-Related Data Exchange: Advanced Technical Support for Interoperability","year":2024,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal; Centre Intégré Universitaire de Santé et de Services Sociaux du Centre-Sud-de-l'Île-de-Montréal","funders":"","keywords":"Interoperability; Computer science; SPARQL; Suite; Ontology; World Wide Web; Web service; Interface (matter); Service (business); Information retrieval; RDF; Semantic Web; Operating system","score_opus":0.049562813651191986,"score_gpt":0.40648886537450196,"score_spread":0.35692605172331,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401824717","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0039133355,0.0002504409,0.9568542,0.0030016631,0.00025773194,0.0006706459,0.0012350072,0.021534054,0.012282928],"genre_scores_gemma":[0.08663221,0.0009647338,0.8789445,0.0021982768,0.00036975025,0.0008339331,0.012299792,0.0052196705,0.012537062],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98039174,0.004561058,0.0029862153,0.0018525731,0.008803558,0.0014048638],"domain_scores_gemma":[0.9556104,0.008079827,0.0014873886,0.02650317,0.0069016037,0.0014175501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03328471,0.0015475589,0.001365989,0.0037568305,0.0022697635,0.009905085,0.0061769076,0.0027813797,0.006653755],"category_scores_gemma":[0.04933356,0.0013109866,0.0020868627,0.004469761,0.0021800138,0.01707949,0.0153744845,0.0048588095,0.0069043906],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008249147,0.0007539252,0.0069984603,0.0012279103,0.00040727708,0.0015022577,0.00349021,0.008407243,0.03202422,0.39191073,0.08770597,0.4647469],"study_design_scores_gemma":[0.00024911872,0.00012322607,0.001965019,0.00049349695,0.00023502762,0.001098962,0.00060586067,0.070623636,0.052874036,0.17805105,0.6934457,0.00023485172],"about_ca_topic_score_codex":0.010599218,"about_ca_topic_score_gemma":0.0067913663,"teacher_disagreement_score":0.03328471,"about_ca_system_score_codex":0.0030602205,"about_ca_system_score_gemma":0.007449998,"threshold_uncertainty_score":0.17602849},"labels":[],"label_agreement":null},{"id":"W4401834311","doi":"10.1101/2024.08.14.24312001","title":"Using Meta-Transformers for Multimodal Clinical Decision Support and Evidence-Based Medicine","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Lakehead University","funders":"","keywords":"Clinical decision making; Transformer; Meta-analysis; Medicine; Computer science; Intensive care medicine; Engineering; Internal medicine; Electrical engineering","score_opus":0.2370704361291506,"score_gpt":0.44886337523489217,"score_spread":0.21179293910574157,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401834311","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09491128,0.011581796,0.7619899,0.01703473,0.0008311382,0.0019615553,0.031157523,0.06292579,0.017606232],"genre_scores_gemma":[0.5335246,0.001795601,0.44720927,0.0014572768,0.00018037474,0.00037942576,0.013154265,0.0008540381,0.0014452034],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9928469,0.003663992,0.0010139145,0.0008651529,0.0013725263,0.0002373926],"domain_scores_gemma":[0.9555352,0.03644774,0.0018277761,0.003032388,0.0023291514,0.0008276772],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015086977,0.0013131466,0.0009945527,0.009586752,0.00065438036,0.005273349,0.0022266675,0.0017512444,0.010963206],"category_scores_gemma":[0.058522247,0.00060613785,0.0024113583,0.0039374414,0.0008488545,0.0068058134,0.004354945,0.0022230612,0.0020648008],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008859918,0.0012091775,0.027833395,0.0064761667,0.0019953116,0.0019742332,0.0019716395,0.0818868,0.013574679,0.090725146,0.039031673,0.724462],"study_design_scores_gemma":[0.0009450317,0.001076937,0.0053938827,0.0024684444,0.001921096,0.0016311167,0.0012044531,0.6088172,0.051605817,0.23416397,0.09042153,0.0003505108],"about_ca_topic_score_codex":0.002412484,"about_ca_topic_score_gemma":0.004661984,"teacher_disagreement_score":0.015086977,"about_ca_system_score_codex":0.0023622462,"about_ca_system_score_gemma":0.003046676,"threshold_uncertainty_score":0.079788506},"labels":[],"label_agreement":null},{"id":"W4401873193","doi":"10.1016/j.apgeog.2024.103392","title":"Spatial intelligence and contextual relevance in AI-driven health information retrieval","year":2024,"lang":"en","type":"article","venue":"Applied Geography","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Relevance (law); Geography; Information retrieval; Spatial analysis; Cartography; Data science; Computer science; Artificial intelligence; Remote sensing; Political science","score_opus":0.00893510160929684,"score_gpt":0.2639046996179468,"score_spread":0.25496959800865,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4401873193","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4634892,0.0015820256,0.5188356,0.0036012675,0.00010634397,0.00038606455,0.000915531,0.0024909377,0.008593001],"genre_scores_gemma":[0.94790095,0.0001675673,0.050585248,0.0001740066,0.00003239496,0.00008064363,0.000315425,0.000062722414,0.0006809757],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9961939,0.0027257586,0.00022892741,0.00042228022,0.00030333915,0.0001258476],"domain_scores_gemma":[0.98101074,0.015735332,0.001138896,0.0008404928,0.0010415133,0.00023304742],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00403525,0.0005204761,0.00044782143,0.0017694446,0.0004065975,0.0020231563,0.0005476685,0.00066412956,0.002417035],"category_scores_gemma":[0.030053861,0.00035100817,0.0005258268,0.0010504199,0.0011881675,0.004241134,0.0015711766,0.0007728983,0.0005166166],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002573898,0.0006369256,0.059226096,0.0023806083,0.00047143575,0.0013362826,0.02037277,0.2146976,0.035089426,0.08891168,0.0059324275,0.56837094],"study_design_scores_gemma":[0.00008504484,0.00047791147,0.016987631,0.00017382269,0.00018433706,0.0005681886,0.004779026,0.8515856,0.009311257,0.10620539,0.0095134415,0.00012832707],"about_ca_topic_score_codex":0.0051771016,"about_ca_topic_score_gemma":0.0043123965,"teacher_disagreement_score":0.0051771016,"about_ca_system_score_codex":0.0014128771,"about_ca_system_score_gemma":0.00088345236,"threshold_uncertainty_score":0.021340668},"labels":[],"label_agreement":null},{"id":"W4402191631","doi":"10.1007/978-981-97-3962-2_21","title":"Ethical Issues in Biomedical Text Mining","year":2024,"lang":"en","type":"book-chapter","venue":"Transactions on computer systems and networks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Engineering ethics; Psychology; Computer science; Data science; Engineering","score_opus":0.0171615787713674,"score_gpt":0.2644848951428124,"score_spread":0.24732331637144503,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402191631","genre_codex":"commentary","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038519946,0.03558512,0.22660501,0.49075553,0.011385554,0.000162122,0.0003318824,0.00020746024,0.23111525],"genre_scores_gemma":[0.30000356,0.04120755,0.24678479,0.14639395,0.038576357,0.001560738,0.0008007883,0.00075351243,0.22391883],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9757642,0.015381816,0.0013322907,0.00092526624,0.0061854958,0.00041096206],"domain_scores_gemma":[0.8814298,0.10465013,0.0019630233,0.0048120567,0.0062456364,0.0008993193],"candidate_categories":["research_integrity"],"consensus_categories":[],"category_scores_codex":[0.032333154,0.00039764578,0.000726999,0.0021462892,0.0025153558,0.007475039,0.0014206371,0.0053882613,0.007684488],"category_scores_gemma":[0.09464897,0.0005280086,0.00036978468,0.002162043,0.011902818,0.011364945,0.0034185103,0.0067103156,0.002555065],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013096268,0.000016625007,0.00018868799,0.00013064589,0.000009170962,0.00011578006,0.0007723531,0.00037863824,0.00011116615,0.8636568,0.07137616,0.06323083],"study_design_scores_gemma":[0.0000067124306,0.000004545324,0.000090527465,0.00017332217,0.000004837826,0.00017419299,0.00025726281,0.0012453286,0.00016467238,0.8980415,0.09983058,0.000006614051],"about_ca_topic_score_codex":0.00077206903,"about_ca_topic_score_gemma":0.0008515611,"teacher_disagreement_score":0.99461174,"about_ca_system_score_codex":0.0020765914,"about_ca_system_score_gemma":0.0036752059,"threshold_uncertainty_score":0.17099607},"labels":[],"label_agreement":null},{"id":"W4402265685","doi":"10.1109/tqcebt59414.2024.10545087","title":"Semantic Reasoning and Knowledge Discovery in Biomedical Informatics Using Domain-Specific Ontologies","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Horizon College and Seminary","funders":"","keywords":"Computer science; Domain (mathematical analysis); Domain knowledge; Open Biomedical Ontologies; Ontology; Knowledge extraction; Data science; Informatics; Semantic Web; Information retrieval; Artificial intelligence; Natural language processing; Semantic Web Stack; OWL-S","score_opus":0.02084493794549056,"score_gpt":0.29637772679623825,"score_spread":0.27553278885074767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402265685","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0038777087,0.0035511442,0.9791084,0.0052457214,0.0001901321,0.00014755852,0.00022098848,0.00033250224,0.0073258914],"genre_scores_gemma":[0.08891977,0.0070300084,0.89778525,0.0013895443,0.00040886085,0.0003143431,0.0009371739,0.00008728817,0.0031277738],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9912089,0.0051502716,0.001040502,0.00070057483,0.0016065731,0.00029324577],"domain_scores_gemma":[0.99190795,0.0059908526,0.00043709736,0.00093817007,0.0005449789,0.00018093231],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011271233,0.0008038427,0.0009767681,0.0059140287,0.0019376932,0.007920991,0.0017867305,0.0018436664,0.0017810591],"category_scores_gemma":[0.011975307,0.000716682,0.0025610742,0.006736903,0.0048834286,0.014535695,0.0045147855,0.003088225,0.00063764775],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031314266,0.000054934873,0.0006385707,0.00050536875,0.00010373683,0.0003551975,0.0013024176,0.0086208945,0.0011617509,0.9032131,0.0038637137,0.080149],"study_design_scores_gemma":[0.000015050889,0.000017624807,0.00033685172,0.00028224388,0.00005752141,0.00028733673,0.00060417,0.050126676,0.0013100368,0.89475787,0.05217171,0.000032980395],"about_ca_topic_score_codex":0.004618556,"about_ca_topic_score_gemma":0.005054518,"teacher_disagreement_score":0.011271233,"about_ca_system_score_codex":0.002575079,"about_ca_system_score_gemma":0.0035917403,"threshold_uncertainty_score":0.059608698},"labels":[],"label_agreement":null},{"id":"W4402405015","doi":"10.23889/ijpds.v9i5.2553","title":"Free Text Analysis: Identification of adverse drug events in clinical notes","year":2024,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba; Manitoba Health","funders":"","keywords":"Identification (biology); Drug; Computer science; Drug reaction; Adverse effect; Natural language processing; Medicine; Pharmacology; Biology","score_opus":0.06472439430834917,"score_gpt":0.44615202051066194,"score_spread":0.38142762620231274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402405015","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7344908,0.004702381,0.14346634,0.003747224,0.00052225066,0.002937163,0.09539331,0.007415066,0.0073254528],"genre_scores_gemma":[0.7580188,0.0013333109,0.178858,0.00065938337,0.0004176277,0.0011564463,0.056403857,0.00020745906,0.002945068],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9943277,0.0020295533,0.0011912562,0.0011632869,0.0011064858,0.00018169347],"domain_scores_gemma":[0.93867964,0.044250008,0.009273827,0.0025793042,0.004586529,0.00063065736],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0048125517,0.0011163135,0.00058759644,0.008603912,0.0005263903,0.0017378333,0.0010150466,0.0010209336,0.0036970307],"category_scores_gemma":[0.037004612,0.00022470113,0.00080700667,0.0035823325,0.0005162693,0.0020858431,0.0011223807,0.0008362672,0.0013149268],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025320998,0.0008033736,0.29524836,0.007413631,0.0006366372,0.00405221,0.0033234425,0.011095027,0.033578966,0.0021243559,0.029681498,0.60951036],"study_design_scores_gemma":[0.0004432559,0.0016323386,0.620885,0.0021474047,0.00069959625,0.007031996,0.0050576124,0.22366747,0.061397467,0.015275658,0.061318632,0.00044359884],"about_ca_topic_score_codex":0.0033626421,"about_ca_topic_score_gemma":0.004712855,"teacher_disagreement_score":0.008603912,"about_ca_system_score_codex":0.00094643066,"about_ca_system_score_gemma":0.0014109362,"threshold_uncertainty_score":0.025451541},"labels":[],"label_agreement":null},{"id":"W4402406304","doi":"10.23889/ijpds.v9i5.2896","title":"Comparing terminology mappings to ICD-10 coded data in Discharge Abstract Database (DAD) in Alberta, Canada","year":2024,"lang":"en","type":"article","venue":"International Journal for Population Data Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Alberta Health Services; University of Calgary","funders":"","keywords":"Terminology; Database; Computer science; Data mining; Linguistics","score_opus":0.0781127730537302,"score_gpt":0.3860818542312556,"score_spread":0.30796908117752536,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402406304","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98523474,0.0010268602,0.00064669695,0.00038905625,0.000025446803,0.00012480642,0.009933412,0.000041305146,0.0025776185],"genre_scores_gemma":[0.9879529,0.00051844615,0.001431305,0.00013444063,0.0000075972835,0.000045315574,0.009379958,0.000012063294,0.0005179361],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99164104,0.0013225189,0.0008295341,0.0010142401,0.0042457115,0.0009470139],"domain_scores_gemma":[0.9734995,0.0048376177,0.00398808,0.0009578213,0.015546717,0.0011704066],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00763367,0.0003348007,0.0003152472,0.0046577356,0.0015430812,0.0021406421,0.0019001983,0.00039151387,0.0011184217],"category_scores_gemma":[0.032002315,0.00032112535,0.00046474324,0.010828986,0.0010365093,0.00052713865,0.0020014113,0.00048578606,0.00019803422],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017461044,0.000024429144,0.985622,0.0001278306,0.000083187864,0.00009660043,0.0018943512,0.00070791726,0.0002719428,0.0003973204,0.0017450664,0.008854749],"study_design_scores_gemma":[0.000009190637,0.000022901371,0.9940905,0.00009111756,0.00003345012,0.000055139877,0.0024321987,0.0012872437,0.00017881171,0.000070317525,0.0017129016,0.000016155964],"about_ca_topic_score_codex":0.98817086,"about_ca_topic_score_gemma":0.9893645,"teacher_disagreement_score":0.03536426,"about_ca_system_score_codex":0.03536426,"about_ca_system_score_gemma":0.033012394,"threshold_uncertainty_score":0.25658685},"labels":[],"label_agreement":null},{"id":"W4402592798","doi":"10.2196/60665","title":"An Automatic and End-to-End System for Rare Disease Knowledge Graph Construction Based on Ontology-Enhanced Large Language Models: Development Study","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Information extraction; Named-entity recognition; Relationship extraction; Domain knowledge; Ontology; Data science; Information retrieval; Data mining; Artificial intelligence","score_opus":0.01387842618304462,"score_gpt":0.3118774236526712,"score_spread":0.29799899746962655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402592798","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033283893,0.00055619515,0.7441306,0.00048640434,0.00013691797,0.0013224324,0.008513264,0.20924112,0.0023291921],"genre_scores_gemma":[0.062239096,0.00035980073,0.90495414,0.00027069906,0.00002732246,0.0006741753,0.02633876,0.0023947493,0.002741359],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985636,0.00032134514,0.00014712897,0.0005599857,0.00032757933,0.00008036095],"domain_scores_gemma":[0.9960639,0.001998824,0.00020985916,0.0006630024,0.00087185483,0.00019257136],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002157309,0.0013321356,0.0007792109,0.002249087,0.00061417336,0.0015595008,0.0023770628,0.0009899603,0.007680943],"category_scores_gemma":[0.006829933,0.00085894275,0.0015146941,0.0013019366,0.00036333973,0.0033554044,0.0024180228,0.0014529736,0.005025571],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00050777174,0.0010462676,0.0065044076,0.0011634318,0.00031559926,0.0010081856,0.00063271675,0.018381061,0.032846034,0.005336033,0.067150034,0.86510843],"study_design_scores_gemma":[0.0003432544,0.000555589,0.006313198,0.00016450076,0.00028643376,0.0016885409,0.00045723186,0.8432567,0.06407448,0.006710287,0.07594008,0.00020978278],"about_ca_topic_score_codex":0.007870907,"about_ca_topic_score_gemma":0.010500312,"teacher_disagreement_score":0.007870907,"about_ca_system_score_codex":0.0009714553,"about_ca_system_score_gemma":0.0027411242,"threshold_uncertainty_score":0.025695324},"labels":[],"label_agreement":null},{"id":"W4402742651","doi":"10.1101/2024.09.20.24314053","title":"A Proof-of-Concept Large Language Model Application to Support Clinical Trial Screening in Surgical Oncology","year":2024,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College","funders":"Genentech; Plexxikon; Regeneron Pharmaceuticals; Ipsen; Pfizer","keywords":"Proof of concept; Clinical trial; Medical physics; Oncology; Medicine; Precision oncology; Internal medicine; Computer science; Cancer","score_opus":0.06516509828343875,"score_gpt":0.4185515456462815,"score_spread":0.35338644736284275,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402742651","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.060253326,0.0002904355,0.8362342,0.0020323396,0.0002974716,0.0015542683,0.0035292273,0.08939327,0.006415501],"genre_scores_gemma":[0.4363439,0.00033887892,0.5500328,0.0010002949,0.00008698661,0.0016485119,0.004137633,0.0024855183,0.00392557],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.998858,0.00050340453,0.00009213869,0.00014862978,0.00031930878,0.00007848922],"domain_scores_gemma":[0.99076617,0.0072604227,0.00036877918,0.0005247088,0.0008014898,0.0002784731],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0045078793,0.0009214986,0.00052436587,0.0005177682,0.0003123182,0.0013636741,0.0020500196,0.0012737504,0.012243022],"category_scores_gemma":[0.017676577,0.0005018588,0.00083175604,0.00026618055,0.0004669188,0.0009697487,0.0013532038,0.001229961,0.0021990852],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020485083,0.0012959881,0.012688409,0.001233044,0.00027948394,0.0018162527,0.0004165067,0.7066115,0.018615533,0.0153613845,0.041831966,0.19780144],"study_design_scores_gemma":[0.00026842684,0.00012990614,0.0003413242,0.00004706468,0.00001916587,0.00007646641,0.0000146807515,0.9852645,0.0036645816,0.0037325567,0.006422145,0.000019197189],"about_ca_topic_score_codex":0.0045637344,"about_ca_topic_score_gemma":0.004334948,"teacher_disagreement_score":0.012243022,"about_ca_system_score_codex":0.0009708133,"about_ca_system_score_gemma":0.002970479,"threshold_uncertainty_score":0.040957034},"labels":[],"label_agreement":null},{"id":"W4402841129","doi":"10.2139/ssrn.4965919","title":"A Computational Framework for Defining and Validating Reproducible Phenotyping Algorithms of 313 Diseases in the UK Biobank","year":2024,"lang":"en","type":"preprint","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Cancer Research","funders":"Medical Research Council","keywords":"Biobank; Computer science; Computational biology; Algorithm; Biology; Bioinformatics","score_opus":0.01860545006647126,"score_gpt":0.316032856469037,"score_spread":0.29742740640256576,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402841129","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008848085,0.00014255253,0.9813176,0.00079385424,0.00003330051,0.0005920681,0.0021764014,0.005327196,0.0007689461],"genre_scores_gemma":[0.077835836,0.00010731334,0.91408235,0.00020816589,0.000021581836,0.0005320089,0.006603944,0.00035272818,0.0002561019],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9767181,0.008500464,0.0049007097,0.0038573914,0.0051457398,0.00087756105],"domain_scores_gemma":[0.93235606,0.041772205,0.004289741,0.013206713,0.0070124674,0.0013627568],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.037174366,0.0012352908,0.0017272917,0.006630366,0.002444896,0.010520902,0.0051202253,0.0028718507,0.0023395526],"category_scores_gemma":[0.098376974,0.0014695304,0.0039190524,0.0051129507,0.0026336785,0.007285945,0.0074380822,0.0032915478,0.00080127764],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066684035,0.00080917706,0.015881618,0.0017391762,0.0011984322,0.0010242925,0.0015879706,0.46588075,0.0075526633,0.28435194,0.019086374,0.20022075],"study_design_scores_gemma":[0.00012318198,0.00010777805,0.0012719308,0.00032451004,0.00020734721,0.00017138306,0.00021893716,0.8842733,0.004186606,0.09728209,0.011779694,0.00005333355],"about_ca_topic_score_codex":0.02137796,"about_ca_topic_score_gemma":0.03040792,"teacher_disagreement_score":0.96282566,"about_ca_system_score_codex":0.00521115,"about_ca_system_score_gemma":0.0129242735,"threshold_uncertainty_score":0.19659913},"labels":[],"label_agreement":null},{"id":"W4402867073","doi":"10.1016/j.cancergen.2024.08.030","title":"28. Addition of non-gene features to the CIViC data model","year":2024,"lang":"en","type":"article","venue":"Cancer Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre","funders":"","keywords":"Computational biology; Computer science; Biology","score_opus":0.040218820368345204,"score_gpt":0.34312883713511233,"score_spread":0.3029100167667671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402867073","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05165056,0.0004382503,0.76968944,0.0053248587,0.0009262748,0.0006382291,0.07191746,0.02762148,0.07179348],"genre_scores_gemma":[0.3455792,0.00041695253,0.53302026,0.0010847216,0.00018785601,0.00051011995,0.07242464,0.0021122347,0.044664092],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991165,0.00015958754,0.00009802776,0.00020942745,0.0003373227,0.00007914181],"domain_scores_gemma":[0.9977586,0.0005126504,0.000082278224,0.0007240547,0.0008123652,0.00011000471],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015842121,0.00034833798,0.00047838895,0.0017716173,0.00073642394,0.0016367151,0.0012682725,0.0008232488,0.01490294],"category_scores_gemma":[0.005220185,0.00025674584,0.0012508621,0.0018207618,0.0004232904,0.001977712,0.0009956297,0.001211734,0.0067233527],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010324079,0.0006194305,0.021295495,0.0005256364,0.00018755098,0.0007393004,0.00026838869,0.047068518,0.0066883573,0.12173312,0.25318545,0.5466563],"study_design_scores_gemma":[0.00017815904,0.00013805645,0.008555697,0.00017451413,0.00016863183,0.0005701728,0.00021519276,0.51053727,0.022726167,0.10403372,0.35261863,0.000083832],"about_ca_topic_score_codex":0.012842566,"about_ca_topic_score_gemma":0.024821052,"teacher_disagreement_score":0.01490294,"about_ca_system_score_codex":0.001154136,"about_ca_system_score_gemma":0.0016509859,"threshold_uncertainty_score":0.04985535},"labels":[],"label_agreement":null},{"id":"W4402910734","doi":"10.1016/j.heliyon.2024.e38448","title":"A framework for integrating biomedical knowledge in Wikidata with open biological and biomedical ontologies and MeSH keywords","year":2024,"lang":"en","type":"article","venue":"Heliyon","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"U.S. National Library of Medicine; Universidade de São Paulo; University of Virginia; Universiteit Maastricht; Wikimedia Foundation; Deanship of Scientific Research, Prince Sattam bin Abdulaziz University; Prince Sattam bin Abdulaziz University","keywords":"Computer science; Open Biomedical Ontologies; Data science; Ontology; World Wide Web; Semantic Web; Epistemology; Ontology alignment; Process ontology","score_opus":0.047712900614528785,"score_gpt":0.3624620320917789,"score_spread":0.31474913147725014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4402910734","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0026405295,0.00021080032,0.9838615,0.00061796414,0.000095337644,0.00047896898,0.002093719,0.0072158193,0.0027853244],"genre_scores_gemma":[0.017388813,0.0002057038,0.97618824,0.00016206517,0.000037778882,0.000381545,0.0042668693,0.00040539293,0.000963605],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9929583,0.0020124996,0.0016005844,0.0014607682,0.001682038,0.00028576923],"domain_scores_gemma":[0.9820494,0.006238378,0.0019971028,0.004874263,0.0039005333,0.0009403845],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.010779609,0.0011648779,0.0010098995,0.016352272,0.002499387,0.0069691297,0.003034674,0.0012304104,0.0018950583],"category_scores_gemma":[0.026061716,0.0010233924,0.0027878324,0.01052842,0.0017103016,0.010514779,0.007655134,0.0024730985,0.0019407847],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016939762,0.00037163316,0.0103355115,0.0024653547,0.00060507166,0.0011544123,0.007005522,0.021320155,0.012847008,0.49750844,0.030071113,0.41614643],"study_design_scores_gemma":[0.000048357022,0.00011852269,0.0040386547,0.0012394441,0.00037645584,0.0010505642,0.00264462,0.14512518,0.016952053,0.33115396,0.49700096,0.00025123064],"about_ca_topic_score_codex":0.014163322,"about_ca_topic_score_gemma":0.023779757,"teacher_disagreement_score":0.99303085,"about_ca_system_score_codex":0.0020430058,"about_ca_system_score_gemma":0.007257969,"threshold_uncertainty_score":0.057008684},"labels":[],"label_agreement":null},{"id":"W4403010754","doi":"10.2196/56955","title":"Disambiguating Clinical Abbreviations by One-to-All Classification: Algorithm Development and Validation Study","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Word2vec; Natural language processing; Artificial intelligence; Word embedding; Context (archaeology); Information retrieval; F1 score; Encoder; Machine learning; Macro; Embedding","score_opus":0.0729430305125116,"score_gpt":0.4031236214268938,"score_spread":0.3301805909143822,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403010754","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58519,0.0111969635,0.3634383,0.0016777456,0.00095844915,0.0022463952,0.0046304585,0.021842463,0.008819191],"genre_scores_gemma":[0.6009951,0.0015152396,0.378089,0.00069498306,0.0001423784,0.0010749055,0.01216299,0.00045184355,0.0048736013],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9975757,0.00080148387,0.00032797075,0.0006203271,0.00044057512,0.0002339542],"domain_scores_gemma":[0.9913202,0.004901385,0.00032999285,0.0008643697,0.0023686432,0.000215359],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005456151,0.0023262934,0.0015809805,0.0023191695,0.00083986117,0.0013315864,0.002469567,0.0021878497,0.0031347896],"category_scores_gemma":[0.013938398,0.0004128294,0.001222656,0.0020710565,0.0005562121,0.001821132,0.0015877745,0.0022661327,0.0021515123],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013678111,0.001304938,0.022789724,0.0005175088,0.00047248005,0.00034688707,0.00023140562,0.10736061,0.005258442,0.00102501,0.018653331,0.84067184],"study_design_scores_gemma":[0.00022856407,0.0006116105,0.004334531,0.00009383292,0.00018236868,0.00036238087,0.0003239609,0.9793682,0.00870512,0.0014571598,0.004295034,0.000037163896],"about_ca_topic_score_codex":0.011442618,"about_ca_topic_score_gemma":0.011168538,"teacher_disagreement_score":0.011442618,"about_ca_system_score_codex":0.0014406334,"about_ca_system_score_gemma":0.002924513,"threshold_uncertainty_score":0.028855205},"labels":[],"label_agreement":null},{"id":"W4403012771","doi":"10.1186/s13326-024-00319-w","title":"MeSH2Matrix: combining MeSH keywords and machine learning for biomedical relation classification based on PubMed","year":2024,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Mila - Quebec Artificial Intelligence Institute","funders":"Craig Newmark Philanthropies; Université de Sousse; Wikimedia Foundation; University of Dayton","keywords":"Computer science; Relation (database); Machine learning; Artificial intelligence; Data science; Information retrieval; Data mining","score_opus":0.029422152292203088,"score_gpt":0.28819216417840865,"score_spread":0.2587700118862056,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403012771","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16443458,0.031676795,0.28467447,0.008058958,0.0025733283,0.0033465254,0.36304098,0.11519584,0.026998486],"genre_scores_gemma":[0.18204786,0.0059599923,0.51283336,0.0013904226,0.0007085139,0.0020693599,0.2870725,0.0012326688,0.006685287],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9969119,0.00057951675,0.00055205775,0.00085548207,0.0009447977,0.00015628155],"domain_scores_gemma":[0.99415475,0.003115826,0.0008222255,0.0008611423,0.00077612576,0.00026997647],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002499224,0.001991443,0.0011991982,0.020038716,0.0011782655,0.0021757993,0.0018669901,0.0014998671,0.0055427398],"category_scores_gemma":[0.014186423,0.0004548721,0.0016581947,0.011479658,0.00044612394,0.0045794887,0.003345016,0.0011309007,0.0044044014],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000859355,0.0005808044,0.049824096,0.005937808,0.0008647975,0.0009343551,0.00064989226,0.00806078,0.015260094,0.0074001914,0.18035755,0.7292703],"study_design_scores_gemma":[0.00050150283,0.0013926887,0.06432586,0.0016119337,0.0009766937,0.0026277394,0.0015896815,0.36059794,0.028035099,0.04556996,0.49237195,0.0003989336],"about_ca_topic_score_codex":0.009602194,"about_ca_topic_score_gemma":0.024058696,"teacher_disagreement_score":0.020038716,"about_ca_system_score_codex":0.0015671528,"about_ca_system_score_gemma":0.0034582152,"threshold_uncertainty_score":0.01909262},"labels":[],"label_agreement":null},{"id":"W4403122659","doi":"10.1038/s41746-024-01267-6","title":"Enabling data linkages for rare diseases in a resilient environment with the SERDIF framework","year":2024,"lang":"en","type":"article","venue":"npj Digital Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trinity College","funders":"Trinity College Dublin; European Commission","keywords":"Usability; Linkage (software); Disease; Climate change; Data science; Environmental health; Medicine; Computer science; Risk analysis (engineering); Environmental resource management; Ecology; Biology; Pathology","score_opus":0.025218402084422484,"score_gpt":0.29574027527185065,"score_spread":0.27052187318742815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403122659","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0150421355,0.00065211445,0.9207412,0.004806334,0.00023506653,0.0005471075,0.013016126,0.038998198,0.005961732],"genre_scores_gemma":[0.123332374,0.00065821427,0.8434055,0.0017808832,0.00011252275,0.00042987344,0.026012313,0.0020502396,0.0022180977],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99289364,0.0026534754,0.0012414193,0.001195969,0.0017133841,0.00030205236],"domain_scores_gemma":[0.9859803,0.0064925635,0.0008838239,0.004017539,0.0019555695,0.0006702529],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.017266655,0.0008436478,0.0007515337,0.004962113,0.0012168225,0.0047679646,0.002156018,0.0013177791,0.004602646],"category_scores_gemma":[0.029457765,0.0006300855,0.002752059,0.0029528122,0.0012245353,0.006939138,0.012142902,0.0017177959,0.0014012947],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011486614,0.0006026222,0.046406504,0.003307736,0.0010851591,0.0028063236,0.006994794,0.052657098,0.013855856,0.32065946,0.110825166,0.43965054],"study_design_scores_gemma":[0.00019333037,0.00022953031,0.0075400723,0.00097111193,0.00030651287,0.0015548856,0.0018091514,0.16458865,0.0112526305,0.2886775,0.5226424,0.00023410807],"about_ca_topic_score_codex":0.007393139,"about_ca_topic_score_gemma":0.011569383,"teacher_disagreement_score":0.017266655,"about_ca_system_score_codex":0.0013757999,"about_ca_system_score_gemma":0.0030155114,"threshold_uncertainty_score":0.091315925},"labels":[],"label_agreement":null},{"id":"W4403458599","doi":"10.1186/s13326-024-00320-3","title":"Dynamic Retrieval Augmented Generation of Ontologies using Artificial Intelligence (DRAGON-AI)","year":2024,"lang":"en","type":"article","venue":"Journal of Biomedical Semantics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":44,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"Basic Energy Sciences; National Center for Advancing Translational Sciences; National Institute of Environmental Health Sciences; National Human Genome Research Institute; Office of Science; Wellcome Trust; Fundação de Amparo à Pesquisa do Estado de São Paulo; U.S. Department of Energy; National Cancer Institute; National Institutes of Health; National Science Foundation","keywords":"Ontology; Computer science; Open Biomedical Ontologies; Domain (mathematical analysis); Information retrieval; Upper ontology; Process ontology; Data science; Ontology-based data integration; Domain knowledge; Artificial intelligence; Suggested Upper Merged Ontology; Natural language processing","score_opus":0.057231127727669855,"score_gpt":0.3550962777765723,"score_spread":0.29786515004890246,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403458599","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05334874,0.0005404789,0.9100574,0.00069799036,0.00018199383,0.0010226289,0.0010351787,0.023353279,0.009762316],"genre_scores_gemma":[0.13061093,0.00021954064,0.86119926,0.00032616963,0.00003600895,0.00042702825,0.003345913,0.0011106336,0.0027245788],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99563414,0.0018499796,0.0003506865,0.0007213807,0.0013048061,0.00013887807],"domain_scores_gemma":[0.98509026,0.008878244,0.00075265695,0.003196223,0.0018533582,0.00022933012],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005706581,0.0010378372,0.00049133674,0.0025885678,0.0009783426,0.0023027346,0.002260288,0.0010190406,0.0032526786],"category_scores_gemma":[0.018339567,0.00043931056,0.0011366442,0.0014070309,0.0011143538,0.0031275798,0.004622564,0.0014790158,0.0010360142],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051599776,0.00069320493,0.006876116,0.0017759732,0.00027014632,0.0009993847,0.004191719,0.04920776,0.041730452,0.045222927,0.039362255,0.8091541],"study_design_scores_gemma":[0.00031005839,0.0004055059,0.003479346,0.0003482081,0.00029690415,0.0014223449,0.0011543275,0.6580945,0.07769476,0.062270403,0.19430058,0.00022299934],"about_ca_topic_score_codex":0.0035493607,"about_ca_topic_score_gemma":0.0064085857,"teacher_disagreement_score":0.005706581,"about_ca_system_score_codex":0.0012092026,"about_ca_system_score_gemma":0.0021414685,"threshold_uncertainty_score":0.03017962},"labels":[],"label_agreement":null},{"id":"W4403997158","doi":"10.47909/ijsmc.137","title":"Health and medical informatics research: Identifying international collaboration patterns at the country and institution level","year":2024,"lang":"en","type":"article","venue":"Iberoamerican Journal of Science Measurement and Communication","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Institution; Health informatics; Informatics; Data science; Geography; Medicine; Political science; Computer science; Nursing; Public health","score_opus":0.1413070740599149,"score_gpt":0.40125482810720164,"score_spread":0.25994775404728676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4403997158","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9445238,0.008137785,0.011047904,0.0014198505,0.00007064285,0.00020330526,0.007189605,0.00013009798,0.027277077],"genre_scores_gemma":[0.9870946,0.0019669705,0.007967253,0.00007893111,0.000060904356,0.00016164858,0.0019978094,0.000020695752,0.00065115246],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9951918,0.0017067526,0.0009705757,0.00074622507,0.001045407,0.00033928314],"domain_scores_gemma":[0.9600867,0.019804591,0.014413107,0.0014230213,0.003175267,0.0010973063],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.006604075,0.00033662916,0.0005551195,0.03837529,0.0007735179,0.004644294,0.00047510353,0.0005468696,0.003617761],"category_scores_gemma":[0.029580688,0.0001445977,0.0006044348,0.06494688,0.0006946606,0.0043158294,0.002815578,0.00032778239,0.0005220904],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000081391285,0.000060605715,0.8884107,0.0016050396,0.0004060996,0.0002723272,0.004423878,0.0012123658,0.0011157175,0.0073701283,0.0018106912,0.093231134],"study_design_scores_gemma":[0.00001708102,0.00012995391,0.9421659,0.001059509,0.00033900712,0.0008263718,0.018504415,0.0044999127,0.0012078056,0.0073149498,0.023893027,0.000042140076],"about_ca_topic_score_codex":0.0023950543,"about_ca_topic_score_gemma":0.0038991282,"teacher_disagreement_score":0.9933959,"about_ca_system_score_codex":0.0010509674,"about_ca_system_score_gemma":0.001836873,"threshold_uncertainty_score":0.034926116},"labels":[],"label_agreement":null},{"id":"W4404399990","doi":"10.48550/arxiv.2411.08010","title":"ExpressivityBench: Can LLMs Communicate Implicitly?","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Atomic Energy of Canada Limited","keywords":"Computer science; Business; Computer security","score_opus":0.08480621804238049,"score_gpt":0.21751195631833528,"score_spread":0.13270573827595478,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404399990","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25022417,0.001210505,0.6951735,0.0036339737,0.00032642225,0.00060031936,0.0071452977,0.018285379,0.02340044],"genre_scores_gemma":[0.84390914,0.00034848732,0.14265715,0.00069263746,0.00007992088,0.00050617446,0.006615594,0.0012562597,0.0039346637],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9879332,0.0075438474,0.0007613644,0.0012927171,0.0021652994,0.00030349617],"domain_scores_gemma":[0.95004356,0.036617804,0.0027012362,0.0071356413,0.002954216,0.00054758467],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008736572,0.0012195122,0.00052320934,0.0012440894,0.00061618997,0.004007959,0.001521862,0.001893324,0.0040548868],"category_scores_gemma":[0.07036897,0.0004332863,0.0006622489,0.0008377084,0.0017078251,0.008347612,0.00334424,0.0022751095,0.0020042926],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002768298,0.0008460491,0.04278572,0.005570173,0.0006636176,0.000979721,0.020227667,0.08919064,0.09460365,0.111089155,0.04569478,0.5855806],"study_design_scores_gemma":[0.00022290251,0.0012895941,0.012497553,0.0007336136,0.0002930996,0.00085858465,0.0054872674,0.6130294,0.089463934,0.17778803,0.09800982,0.0003261531],"about_ca_topic_score_codex":0.0015900187,"about_ca_topic_score_gemma":0.0017041903,"teacher_disagreement_score":0.008736572,"about_ca_system_score_codex":0.0010142825,"about_ca_system_score_gemma":0.00092660857,"threshold_uncertainty_score":0.04620391},"labels":[],"label_agreement":null},{"id":"W4404610709","doi":"10.1007/978-3-031-77792-9_23","title":"Ontology-Constrained Generation of Domain-Specific Clinical Summaries","year":2024,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal; Héma-Québec","funders":"","keywords":"Computer science; Ontology; Domain (mathematical analysis); Information retrieval; Natural language processing; Artificial intelligence","score_opus":0.04570157674357549,"score_gpt":0.3071494366296804,"score_spread":0.2614478598861049,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404610709","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05462617,0.0008666705,0.8788453,0.0011378423,0.00032838993,0.0010844555,0.026568891,0.029205505,0.0073368135],"genre_scores_gemma":[0.23082612,0.00055365363,0.71057916,0.0002945674,0.0000990104,0.00050547725,0.05271621,0.0010857022,0.0033400818],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99932384,0.00015457014,0.00008758903,0.00018709678,0.00019386293,0.000053151787],"domain_scores_gemma":[0.99741673,0.0014168505,0.00016460984,0.00030476297,0.0005958224,0.00010121211],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010406283,0.000809617,0.0006717181,0.0032841552,0.0004833604,0.0013093918,0.0011479044,0.0008172141,0.0068011703],"category_scores_gemma":[0.006452267,0.00042091694,0.0013825268,0.0022481156,0.00026907993,0.000954081,0.0017254378,0.0007661066,0.0018149578],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078234926,0.00040898085,0.005396846,0.0012804319,0.00030415578,0.0015267959,0.000610445,0.06876522,0.024969172,0.018014476,0.08051915,0.797422],"study_design_scores_gemma":[0.00029115935,0.00017597975,0.0039324043,0.00022071225,0.00031930336,0.0009166735,0.0005744882,0.84817195,0.040592935,0.043756455,0.06096713,0.000080787446],"about_ca_topic_score_codex":0.0054163393,"about_ca_topic_score_gemma":0.011009494,"teacher_disagreement_score":0.0068011703,"about_ca_system_score_codex":0.00087070005,"about_ca_system_score_gemma":0.0026221517,"threshold_uncertainty_score":0.022752166},"labels":[],"label_agreement":null},{"id":"W4404619910","doi":"10.3897/biss.8.142382","title":"Relation Extraction From Unstructured Species Descriptions Using TaxonNERD and LLaMA 2 7B","year":2024,"lang":"en","type":"article","venue":"Biodiversity Information Science and Standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Ministerio de Ciencia Tecnología y Telecomunicaciones; Instituto Tecnológico de Costa Rica; Consejo Superior Universitario Centroamericano; International Development Research Centre","keywords":"Relationship extraction; Adaptability; Computer science; Relation (database); Biodiversity; Taxonomy (biology); False positive paradox; Trophic level; Artificial intelligence; Natural language processing; Information extraction; Information retrieval; Ecology; Biology; Data mining","score_opus":0.02713421106529408,"score_gpt":0.27947845760268597,"score_spread":0.2523442465373919,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404619910","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040619515,0.0034053759,0.41696694,0.001472415,0.0005156109,0.0016419007,0.26346177,0.2606799,0.011236541],"genre_scores_gemma":[0.054975703,0.0006668357,0.6786078,0.00043799164,0.000041726467,0.0014403797,0.25721568,0.0032913552,0.0033224535],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99830675,0.0002472202,0.00041365655,0.0005791762,0.00034974478,0.000103401166],"domain_scores_gemma":[0.9977437,0.00107273,0.0002673504,0.0004656071,0.0003602216,0.00009028926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017799591,0.0024838177,0.0013034479,0.010129886,0.0012332675,0.0035575884,0.0018158089,0.0017087304,0.012104195],"category_scores_gemma":[0.0078328205,0.0013189028,0.004142971,0.0050931186,0.00056465855,0.005137239,0.003547941,0.001996156,0.007533046],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012344164,0.0006868558,0.032240663,0.008314868,0.0009489989,0.0034573118,0.0051967287,0.017361203,0.038374647,0.031052327,0.23465699,0.626475],"study_design_scores_gemma":[0.00029127274,0.00019751872,0.016706422,0.0009622631,0.00041155407,0.0019189969,0.001870322,0.31634355,0.031358022,0.02888452,0.6007677,0.00028783924],"about_ca_topic_score_codex":0.01804214,"about_ca_topic_score_gemma":0.034987826,"teacher_disagreement_score":0.01804214,"about_ca_system_score_codex":0.0020932746,"about_ca_system_score_gemma":0.0028005608,"threshold_uncertainty_score":0.040492535},"labels":[],"label_agreement":null},{"id":"W4404620563","doi":"10.1186/s13643-024-02699-7","title":"Computer-assisted screening in systematic evidence synthesis requires robust and well-evaluated stopping criteria","year":2024,"lang":"en","type":"letter","venue":"Systematic Reviews","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"European Commission; Wellcome Trust","keywords":"Medicine; Systematic review; Evidence-based medicine; Medical physics; MEDLINE; Alternative medicine; Pathology","score_opus":0.128503148657038,"score_gpt":0.3552626784208391,"score_spread":0.2267595297638011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404620563","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":"methods","domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":"methods","domain_consensus":"methods","prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.004244949,0.018878743,0.16918845,0.7639266,0.023079315,0.0060545662,0.0016645755,0.0012340477,0.011728742],"genre_scores_gemma":[0.088504404,0.008276119,0.46451774,0.39313686,0.020180346,0.020935277,0.0007292249,0.00049081515,0.0032291568],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.32859653,0.5348394,0.09300765,0.009817093,0.03171685,0.0020225001],"domain_scores_gemma":[0.023068687,0.9373601,0.014206859,0.009007933,0.014727874,0.00162851],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.509275,0.0015404167,0.0071177036,0.0058641317,0.0026536477,0.012029949,0.0050521283,0.03013536,0.008487554],"category_scores_gemma":[0.878509,0.002680857,0.0076974793,0.006468593,0.006335808,0.010257844,0.0054721492,0.027168678,0.004570891],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0070828646,0.00042253177,0.0057027135,0.027196454,0.006227027,0.0019972308,0.002840132,0.006955229,0.0009675983,0.09430168,0.45713848,0.38916802],"study_design_scores_gemma":[0.008702395,0.0021221247,0.0053820093,0.043372996,0.0042773844,0.0020076053,0.00061541284,0.05346973,0.0027912045,0.66467386,0.2115024,0.0010828812],"about_ca_topic_score_codex":0.0021202748,"about_ca_topic_score_gemma":0.003662593,"teacher_disagreement_score":0.49072498,"about_ca_system_score_codex":0.0076768002,"about_ca_system_score_gemma":0.020747943,"threshold_uncertainty_score":0.6051513},"labels":[],"label_agreement":null},{"id":"W4404667492","doi":"10.2196/60095","title":"Developing an ICD-10 Coding Assistant: Pilot Study Using RoBERTa and GPT-4 for Term Extraction and Description-Based Code Selection","year":2024,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Coding (social sciences); Selection (genetic algorithm); Code (set theory); Term (time); Computer science; Statistics; Mathematics; Artificial intelligence; World Wide Web; Programming language","score_opus":0.18300896160671315,"score_gpt":0.46883495622108085,"score_spread":0.2858259946143677,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404667492","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7812924,0.000723536,0.15145274,0.0017971761,0.00060952833,0.004699996,0.024146741,0.029016908,0.0062611005],"genre_scores_gemma":[0.5358311,0.0004381171,0.37227082,0.0008691289,0.00015934302,0.0035988423,0.07664649,0.0014729508,0.008713236],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970837,0.0013792585,0.0002445301,0.00074799097,0.00038007475,0.00016435115],"domain_scores_gemma":[0.9859119,0.008563716,0.0003871259,0.0017950584,0.002696585,0.0006455979],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0076713325,0.0009101961,0.000587051,0.0013876415,0.00057075836,0.000988553,0.001642342,0.0008745989,0.006516016],"category_scores_gemma":[0.021519057,0.0003241536,0.0007568752,0.0011685827,0.00047608907,0.0013095544,0.0019575981,0.0017813687,0.0034755494],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023178954,0.003316548,0.036878668,0.0015807623,0.00021173785,0.0023055132,0.005655663,0.021582425,0.03173775,0.0020553959,0.085861996,0.80649567],"study_design_scores_gemma":[0.0037954245,0.0055144317,0.10082682,0.00054219353,0.0005416414,0.0030337344,0.009341121,0.5419566,0.12699911,0.0062126527,0.20074765,0.00048851455],"about_ca_topic_score_codex":0.012980488,"about_ca_topic_score_gemma":0.014226036,"teacher_disagreement_score":0.012980488,"about_ca_system_score_codex":0.0014472606,"about_ca_system_score_gemma":0.0026035903,"threshold_uncertainty_score":0.04057032},"labels":[],"label_agreement":null},{"id":"W4404715768","doi":"10.1080/14796708.2024.2419271","title":"Focusing on earlier diagnosis of Alzheimer's disease: a plain language summary","year":2024,"lang":"en","type":"article","venue":"Future Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Novo Nordisk","keywords":"Plain language; Disease; Plain English; Medicine; Alzheimer's disease; Neurology; Psychology; Psychiatry; Linguistics; Pathology; Philosophy","score_opus":0.009494694994666247,"score_gpt":0.2667872587512092,"score_spread":0.2572925637565429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404715768","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021342512,0.29287618,0.012537286,0.48696142,0.15739019,0.0007865719,0.01225733,0.0009366623,0.03412009],"genre_scores_gemma":[0.018782778,0.4477857,0.0308665,0.2879676,0.15820411,0.0014191045,0.01728083,0.00093080633,0.036762644],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99509203,0.0016342516,0.0018113244,0.00036042812,0.0008911371,0.0002109452],"domain_scores_gemma":[0.9574256,0.01740136,0.00461961,0.0006984552,0.018243019,0.0016118911],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0059461603,0.0016306753,0.0018321867,0.0052186255,0.0009894874,0.004804643,0.0014674237,0.0028932213,0.028294606],"category_scores_gemma":[0.038733672,0.00049272686,0.0023095158,0.0031098952,0.00084930187,0.0067169797,0.0016509168,0.0047650198,0.013073326],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002277695,0.000043317286,0.0008746445,0.009604964,0.00011893928,0.00033061224,0.00033265122,0.00020088919,0.0006198105,0.0026753645,0.8575721,0.12739898],"study_design_scores_gemma":[0.000038815848,0.00010978169,0.0025584,0.012208499,0.00017694826,0.00070236,0.00036205823,0.00013296616,0.0003514869,0.0022342177,0.9810681,0.00005637276],"about_ca_topic_score_codex":0.0055794097,"about_ca_topic_score_gemma":0.007553236,"teacher_disagreement_score":0.028294606,"about_ca_system_score_codex":0.0027796347,"about_ca_system_score_gemma":0.005552099,"threshold_uncertainty_score":0.09465486},"labels":[],"label_agreement":null},{"id":"W4404781320","doi":"10.18653/v1/2024.tsar-1.5","title":"Cochrane-auto: An Aligned Dataset for the Simplification of Biomedical Abstracts","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Nederlandse Organisatie voor Wetenschappelijk Onderzoek; Universiteit van Amsterdam; Canadian Institute of Steel Construction","keywords":"Computer science; Information retrieval; Data science","score_opus":0.030635199986428538,"score_gpt":0.3629747871951487,"score_spread":0.33233958720872014,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404781320","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.027244305,0.0029667777,0.02155839,0.0012292502,0.0006567048,0.0011723135,0.91561157,0.022323929,0.0072367443],"genre_scores_gemma":[0.012739556,0.00038045942,0.038920186,0.00022431741,0.000084966574,0.0010864193,0.94424206,0.0006779228,0.001644175],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968851,0.0008835563,0.00064348296,0.00076178793,0.0007222326,0.00010385258],"domain_scores_gemma":[0.9888045,0.0059396764,0.00095020176,0.0018200246,0.0018624233,0.0006232381],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029425258,0.001851899,0.00086516736,0.007829586,0.0011817192,0.0015044312,0.002022012,0.0023541902,0.012414392],"category_scores_gemma":[0.02082842,0.00067320047,0.0012016967,0.0039137583,0.0006049561,0.002356327,0.0029838134,0.0018912419,0.01023089],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00090180116,0.0005095323,0.005608863,0.0096333595,0.00033273018,0.0013516353,0.0011188153,0.0032153532,0.017678965,0.0032582784,0.85307544,0.10331531],"study_design_scores_gemma":[0.0009939155,0.00030928323,0.019741239,0.0007894684,0.00028966906,0.0016729902,0.0009280095,0.024107318,0.02247689,0.0051179244,0.92330754,0.00026579853],"about_ca_topic_score_codex":0.0069714305,"about_ca_topic_score_gemma":0.017729338,"teacher_disagreement_score":0.012414392,"about_ca_system_score_codex":0.0010936831,"about_ca_system_score_gemma":0.003376643,"threshold_uncertainty_score":0.04153031},"labels":[],"label_agreement":null},{"id":"W4404788982","doi":"10.32920/27921843.v1","title":"Measuring Knowledge Translation Uptake Using Citation Metrics: A Case Study of a Pan-Canadian Network of Pharmacoepidemiology Researchers","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Pharmacoepidemiology; Citation; Knowledge translation; Translation (biology); Computer science; Data science; Business; Knowledge management; Medicine; Chemistry; Library science; Pharmacology; Biochemistry","score_opus":0.46919522640668426,"score_gpt":0.4489153994229413,"score_spread":0.020279826983742977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404788982","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89766973,0.0046835626,0.003676935,0.053286552,0.00025508166,0.00093784503,0.0025150147,0.00012400594,0.036851257],"genre_scores_gemma":[0.9840731,0.002478147,0.005751735,0.0032756492,0.000093843955,0.00030665344,0.00069128524,0.000077050136,0.0032525891],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9499379,0.020870851,0.004006348,0.0036104326,0.016337281,0.0052371942],"domain_scores_gemma":[0.73611116,0.1322644,0.018683894,0.010860595,0.084423326,0.017656563],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.064556845,0.0005624608,0.0012245546,0.010161002,0.020696571,0.010296689,0.0038224168,0.0028579566,0.003350616],"category_scores_gemma":[0.21451646,0.00050939317,0.000642661,0.029674726,0.0041666655,0.0070623443,0.0068627964,0.0021426324,0.00052617234],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003250211,0.000434445,0.3028954,0.0020821625,0.00023629366,0.005420771,0.47162774,0.00082637963,0.00094879937,0.012374205,0.027769277,0.17505944],"study_design_scores_gemma":[0.00015541191,0.0003536265,0.25887388,0.0033316857,0.0004413876,0.002999178,0.56587994,0.0052701253,0.0016519179,0.006827966,0.15383692,0.0003778974],"about_ca_topic_score_codex":0.8785495,"about_ca_topic_score_gemma":0.894227,"teacher_disagreement_score":0.989839,"about_ca_system_score_codex":0.06674178,"about_ca_system_score_gemma":0.11289064,"threshold_uncertainty_score":0.48424774},"labels":[],"label_agreement":null},{"id":"W4404789135","doi":"10.32920/27921843","title":"Measuring Knowledge Translation Uptake Using Citation Metrics: A Case Study of a Pan-Canadian Network of Pharmacoepidemiology Researchers","year":2024,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Pharmacoepidemiology; Citation; Knowledge translation; Translation (biology); Knowledge management; Business; Data science; Computer science; Medicine; Chemistry; Library science; Pharmacology; Biochemistry","score_opus":0.46919522640668426,"score_gpt":0.4489153994229413,"score_spread":0.020279826983742977,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404789135","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89766973,0.0046835626,0.003676935,0.053286552,0.00025508166,0.00093784503,0.0025150147,0.00012400594,0.036851257],"genre_scores_gemma":[0.9840731,0.002478147,0.005751735,0.0032756492,0.000093843955,0.00030665344,0.00069128524,0.000077050136,0.0032525891],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9499379,0.020870851,0.004006348,0.0036104326,0.016337281,0.0052371942],"domain_scores_gemma":[0.73611116,0.1322644,0.018683894,0.010860595,0.084423326,0.017656563],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.064556845,0.0005624608,0.0012245546,0.010161002,0.020696571,0.010296689,0.0038224168,0.0028579566,0.003350616],"category_scores_gemma":[0.21451646,0.00050939317,0.000642661,0.029674726,0.0041666655,0.0070623443,0.0068627964,0.0021426324,0.00052617234],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003250211,0.000434445,0.3028954,0.0020821625,0.00023629366,0.005420771,0.47162774,0.00082637963,0.00094879937,0.012374205,0.027769277,0.17505944],"study_design_scores_gemma":[0.00015541191,0.0003536265,0.25887388,0.0033316857,0.0004413876,0.002999178,0.56587994,0.0052701253,0.0016519179,0.006827966,0.15383692,0.0003778974],"about_ca_topic_score_codex":0.8785495,"about_ca_topic_score_gemma":0.894227,"teacher_disagreement_score":0.989839,"about_ca_system_score_codex":0.06674178,"about_ca_system_score_gemma":0.11289064,"threshold_uncertainty_score":0.48424774},"labels":[],"label_agreement":null},{"id":"W4405015797","doi":"10.2196/66088","title":"Microorganisms Linked to Health Care–Associated Infections: Modernization of Terminology Resources for Reporting to the National Healthcare Safety Network","year":2024,"lang":"en","type":"article","venue":"JMIR Public Health and Surveillance","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Health care; Terminology; Medicine; Environmental health; Business; Medical emergency; Computer science; World Wide Web; Political science","score_opus":0.04031266660753761,"score_gpt":0.35378538040203156,"score_spread":0.31347271379449393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405015797","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22811796,0.006456303,0.58198047,0.092947714,0.0022718264,0.0023048292,0.0019955765,0.013490694,0.07043471],"genre_scores_gemma":[0.24280484,0.0027930085,0.73973185,0.004137893,0.00046765213,0.00041426191,0.0031973482,0.0011350544,0.005318077],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9816939,0.0068225404,0.0031135313,0.001212837,0.006500454,0.0006566303],"domain_scores_gemma":[0.93740135,0.020664407,0.008622836,0.019754386,0.012159214,0.0013977757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029131053,0.0006372292,0.00030218528,0.003154336,0.0018429023,0.0065298546,0.003413371,0.0011635165,0.002850605],"category_scores_gemma":[0.056652743,0.00048811367,0.0009679182,0.0035578131,0.002716548,0.008927763,0.006591052,0.0032465744,0.0011729437],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018416755,0.0005230362,0.051252678,0.001481533,0.00009390718,0.0029684016,0.013775097,0.0045019127,0.023678066,0.03481975,0.045814548,0.82090694],"study_design_scores_gemma":[0.000057791152,0.00040399694,0.03880665,0.003875405,0.00022020366,0.0089519285,0.009781613,0.033816934,0.041987624,0.026169378,0.83564115,0.00028728406],"about_ca_topic_score_codex":0.013618139,"about_ca_topic_score_gemma":0.014494906,"teacher_disagreement_score":0.029131053,"about_ca_system_score_codex":0.004582907,"about_ca_system_score_gemma":0.010514928,"threshold_uncertainty_score":0.15406156},"labels":[],"label_agreement":null},{"id":"W4405095249","doi":"10.1016/j.expneurol.2024.115100","title":"Data reporting quality and semantic interoperability increase with community-based data elements (CoDEs). Analysis of the open data commons for spinal cord injury (ODC-SCI)","year":2024,"lang":"en","type":"article","venue":"Experimental Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Research Institute for Aging; Women and Children’s Health Research Institute","funders":"National Institute of Neurological Disorders and Stroke; Division of Mathematical Sciences; National Institutes of Health; Canada Excellence Research Chairs, Government of Canada; University of California; Canada Research Chairs; Craig H. Neilsen Foundation; Wings for Life; Canadian Institutes of Health Research; U.S. Department of Veterans Affairs","keywords":"Interoperability; Spinal cord injury; Commons; Computer science; Data quality; Quality (philosophy); Data mining; Spinal cord; Psychology; Neuroscience; World Wide Web; Engineering; Political science; Operations management","score_opus":0.27932022289204245,"score_gpt":0.48586138777278703,"score_spread":0.2065411648807446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405095249","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2148243,0.010695947,0.6252866,0.06039284,0.002124694,0.0065615824,0.026536293,0.003937522,0.049640205],"genre_scores_gemma":[0.5440841,0.003365289,0.40982193,0.009590309,0.0010205716,0.0074537448,0.021194126,0.0013266951,0.002143201],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.37415695,0.3258444,0.108475216,0.038256723,0.14942026,0.0038463802],"domain_scores_gemma":[0.07277675,0.60390925,0.09327244,0.1405421,0.08693679,0.0025626698],"candidate_categories":["metaresearch","open_science"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.42409724,0.001110505,0.002111778,0.027658328,0.0046211802,0.015192593,0.006092986,0.004149086,0.0033584118],"category_scores_gemma":[0.7376728,0.001887429,0.0033285331,0.03036849,0.01413829,0.024756547,0.022307683,0.0071539804,0.0011605727],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000629397,0.0003647318,0.3204926,0.015542142,0.0016446084,0.0004958356,0.06037665,0.0054893764,0.005180625,0.16135769,0.032496557,0.39592987],"study_design_scores_gemma":[0.00025455415,0.00064099714,0.23677917,0.021584814,0.0010791121,0.0016975051,0.035129942,0.012905376,0.01532226,0.30877554,0.36459795,0.0012327831],"about_ca_topic_score_codex":0.009821083,"about_ca_topic_score_gemma":0.008563106,"teacher_disagreement_score":0.99390703,"about_ca_system_score_codex":0.012011594,"about_ca_system_score_gemma":0.023317002,"threshold_uncertainty_score":0.71019065},"labels":[],"label_agreement":null},{"id":"W4405107639","doi":"10.1016/b978-0-443-15568-0.00010-8","title":"Benefits and challenges of OMICS data integration at the pathway level","year":2024,"lang":"en","type":"book-chapter","venue":"Elsevier eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Omics; Computational biology; Computer science; Data science; Biology; Bioinformatics","score_opus":0.09634372172298897,"score_gpt":0.27706026390770994,"score_spread":0.18071654218472097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405107639","genre_codex":"methods","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010741708,0.16859546,0.5970258,0.08544255,0.0054049953,0.0001658884,0.005966967,0.00378668,0.12286991],"genre_scores_gemma":[0.03779859,0.12067816,0.7778492,0.013216543,0.0041443603,0.00021961644,0.007873834,0.0016350163,0.0365847],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9963264,0.000851089,0.00040053198,0.00041584123,0.0018643057,0.00014180868],"domain_scores_gemma":[0.98406607,0.0117938,0.0003688356,0.0018771933,0.0014701933,0.00042390797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0118714925,0.0012098593,0.0015406125,0.005753988,0.0009901682,0.01408433,0.0025282046,0.002048314,0.008442353],"category_scores_gemma":[0.01086628,0.00097620056,0.0018155852,0.010080732,0.0029484495,0.019495627,0.0051909685,0.0049089165,0.0056824116],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000098827775,0.00008948249,0.0016571505,0.0018037569,0.00018965515,0.00032025465,0.00077793095,0.002433883,0.0047119525,0.295266,0.048537657,0.6441134],"study_design_scores_gemma":[0.0000112034295,0.000034131874,0.0010831128,0.0011063227,0.00010892606,0.0006936111,0.0005722463,0.007049785,0.0026004429,0.447288,0.5393945,0.000057738143],"about_ca_topic_score_codex":0.0023602173,"about_ca_topic_score_gemma":0.0029664594,"teacher_disagreement_score":0.01408433,"about_ca_system_score_codex":0.0013978007,"about_ca_system_score_gemma":0.0025031655,"threshold_uncertainty_score":0.06278318},"labels":[],"label_agreement":null},{"id":"W4405750175","doi":"10.1016/j.jval.2024.10.2510","title":"MT40 Systematic Review Tools Integrated with Artificial Intelligence for Data Extraction: Feature Analysis","year":2024,"lang":"en","type":"article","venue":"Value in Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Micropharma (Canada)","funders":"","keywords":"Computer science; Data extraction; Artificial intelligence; Data mining; Pattern recognition (psychology); Chemistry; MEDLINE","score_opus":0.14167520564103186,"score_gpt":0.40468179563162077,"score_spread":0.2630065899905889,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405750175","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.025879493,0.20751505,0.100963175,0.006901428,0.0010182257,0.019427806,0.59784734,0.024018547,0.016428897],"genre_scores_gemma":[0.12980829,0.10912787,0.45487252,0.0033825694,0.0007022388,0.07106968,0.21507764,0.003198731,0.012760532],"study_design_codex":"systematic_review","study_design_gemma":"observational","domain_scores_codex":[0.9890428,0.0034135787,0.0049191946,0.001054701,0.0013134708,0.00025628283],"domain_scores_gemma":[0.94760996,0.038985953,0.0063303076,0.002307581,0.0040615615,0.0007046932],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.014420291,0.0017397791,0.004299165,0.030901011,0.0013155888,0.0036185395,0.0017472577,0.0010113472,0.033317406],"category_scores_gemma":[0.06188867,0.0009041275,0.0065567596,0.021191834,0.0005693932,0.002111616,0.003867719,0.0008349574,0.0040803794],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014785653,0.00009429005,0.0065034856,0.57713884,0.012186427,0.0007023383,0.0014263901,0.0012810611,0.005576696,0.0056598936,0.08128942,0.30666268],"study_design_scores_gemma":[0.0020741147,0.00083064695,0.027634332,0.18087798,0.066647716,0.0015668868,0.0010566853,0.004583287,0.010474003,0.015805239,0.6879713,0.00047788603],"about_ca_topic_score_codex":0.004268425,"about_ca_topic_score_gemma":0.013662059,"teacher_disagreement_score":0.9855797,"about_ca_system_score_codex":0.0019517286,"about_ca_system_score_gemma":0.015267234,"threshold_uncertainty_score":0.111457884},"labels":[],"label_agreement":null},{"id":"W4405766324","doi":"10.48550/arxiv.2412.16701","title":"AlzheimerRAG: Multimodal Retrieval Augmented Generation for Clinical Use Cases using PubMed articles","year":2024,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Information retrieval; Multimodal therapy; Computer science; Psychology; Psychotherapist","score_opus":0.4198376361877568,"score_gpt":0.3135911284571179,"score_spread":0.10624650773063893,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405766324","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21528068,0.004482207,0.5322959,0.0039960677,0.00057985313,0.0038826927,0.032487523,0.18598497,0.021010222],"genre_scores_gemma":[0.39177606,0.0010893951,0.5758212,0.0010880516,0.00015551619,0.0010377752,0.021919103,0.0015491322,0.0055637415],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984118,0.0007351767,0.00018337418,0.00026517373,0.00034647965,0.000057933274],"domain_scores_gemma":[0.9938484,0.0046725823,0.00029246026,0.00064368424,0.000373832,0.00016901155],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025312002,0.0013418422,0.00044113136,0.0036671702,0.00036467696,0.0013301882,0.0011912766,0.001435129,0.012689293],"category_scores_gemma":[0.012229801,0.00031306947,0.0009881661,0.0013045992,0.00045326803,0.0015213737,0.0022656452,0.0006014879,0.0032130936],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016600412,0.0007251609,0.0092657935,0.003091402,0.00045640115,0.0038991256,0.0020803607,0.021707248,0.039610557,0.004447528,0.07870078,0.83435553],"study_design_scores_gemma":[0.0014686944,0.0018179205,0.01372711,0.000614687,0.00061288057,0.0063323746,0.0030667158,0.64540565,0.093627624,0.036765948,0.1962358,0.00032458772],"about_ca_topic_score_codex":0.0023413585,"about_ca_topic_score_gemma":0.004517249,"teacher_disagreement_score":0.012689293,"about_ca_system_score_codex":0.0005872116,"about_ca_system_score_gemma":0.00072324876,"threshold_uncertainty_score":0.04244989},"labels":[],"label_agreement":null},{"id":"W4405776069","doi":"10.2196/67984","title":"A Validation Tool (VaPCE) for Postcoordinated SNOMED CT Expressions: Development and Usability Study","year":2024,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"SNOMED CT; Computer science; Interoperability; Systematized Nomenclature of Medicine; Semantic interoperability; Correctness; Terminology; Unified Medical Language System; Health informatics; Information retrieval; Software engineering; Data mining; World Wide Web; Health care; Programming language","score_opus":0.02329890040067492,"score_gpt":0.33315772531294885,"score_spread":0.3098588249122739,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405776069","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53883386,0.0027603405,0.3993407,0.0013056601,0.00063917035,0.013092213,0.0049303235,0.029338459,0.009759289],"genre_scores_gemma":[0.35641304,0.0011939066,0.6154559,0.00069265923,0.00008288041,0.009074789,0.006957378,0.0063873567,0.0037422082],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9785824,0.012909368,0.0025803447,0.0017904628,0.0035531279,0.00058429135],"domain_scores_gemma":[0.8554527,0.12090071,0.0023945118,0.005958201,0.014198512,0.0010952853],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.045532458,0.0020360074,0.0012468357,0.004649888,0.001111359,0.002805155,0.0028994896,0.0018871498,0.0040638414],"category_scores_gemma":[0.11725971,0.0011142556,0.001833332,0.0021552325,0.0015044414,0.0045417505,0.0046982383,0.0018183101,0.0015720108],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0027516594,0.004017494,0.035147004,0.011537305,0.00047953764,0.00408318,0.057883456,0.0073184907,0.055179775,0.007862355,0.04297013,0.7707696],"study_design_scores_gemma":[0.0045884047,0.012661841,0.10525723,0.020530393,0.0020054183,0.015485969,0.03610273,0.16132441,0.14427337,0.016664648,0.47893625,0.0021693052],"about_ca_topic_score_codex":0.0018189083,"about_ca_topic_score_gemma":0.0021022893,"teacher_disagreement_score":0.045532458,"about_ca_system_score_codex":0.0012109454,"about_ca_system_score_gemma":0.0032993774,"threshold_uncertainty_score":0.24080151},"labels":[],"label_agreement":null},{"id":"W4405778859","doi":"10.1109/mce.2024.3522521","title":"MedVLM: Medical Vision–Language Model for Consumer Devices","year":2024,"lang":"en","type":"article","venue":"IEEE Consumer Electronics Magazine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Computer science","score_opus":0.014010590841138752,"score_gpt":0.3208812965795662,"score_spread":0.30687070573842745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405778859","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013274476,0.0012154728,0.95872337,0.0012224705,0.00023479255,0.0002315284,0.0037311097,0.017938579,0.0034282566],"genre_scores_gemma":[0.3463407,0.0009953859,0.62443984,0.0020003454,0.00018021093,0.0009994213,0.013026739,0.0014550819,0.010562338],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99951816,0.00014407725,0.00003458827,0.00013919546,0.000121946854,0.00004212249],"domain_scores_gemma":[0.999463,0.00032631564,0.000036576388,0.00006543595,0.00008027609,0.000028413333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010366811,0.00077773206,0.0005371225,0.00087434484,0.00027278534,0.0011775139,0.0017873811,0.0012927189,0.005031273],"category_scores_gemma":[0.0038252324,0.00042209824,0.0016554956,0.0004854998,0.00036275777,0.0011244371,0.0012971356,0.0014293691,0.00215944],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041758682,0.00023544936,0.002972039,0.00043727783,0.00021458587,0.00041376872,0.00026243116,0.4661816,0.008294191,0.025404014,0.06511719,0.4300499],"study_design_scores_gemma":[0.000032541724,0.00004037654,0.00021374311,0.000021704045,0.000012944747,0.00008714109,0.000016853404,0.9776118,0.0013596353,0.011550696,0.009036947,0.000015667314],"about_ca_topic_score_codex":0.011746283,"about_ca_topic_score_gemma":0.015716134,"teacher_disagreement_score":0.011746283,"about_ca_system_score_codex":0.0011560891,"about_ca_system_score_gemma":0.0012918377,"threshold_uncertainty_score":0.023355842},"labels":[],"label_agreement":null},{"id":"W4405868468","doi":"10.1145/3696409.3700264","title":"Development of a Chinese Synonym Library: Enhancing Clinical Terminology Standardization and Interoperability","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Standardization; Terminology; Interoperability; Synonym (taxonomy); Computer science; SNOMED CT; Information retrieval; World Wide Web; Linguistics","score_opus":0.018403178702685704,"score_gpt":0.33259270790554785,"score_spread":0.31418952920286214,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405868468","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042211436,0.0010232101,0.91878045,0.002662518,0.00043878623,0.0021143046,0.008727963,0.013628849,0.010412416],"genre_scores_gemma":[0.08391737,0.0007014761,0.88921046,0.0005729619,0.00011742212,0.00094620883,0.019912438,0.001274805,0.0033469547],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9912361,0.002570775,0.0024905107,0.001316671,0.0020709115,0.00031495883],"domain_scores_gemma":[0.9737429,0.0059583196,0.0016328807,0.0050011678,0.012489292,0.0011753788],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013109737,0.0010987745,0.0015225514,0.012619232,0.0024641696,0.004251051,0.0023137871,0.0009438049,0.005341482],"category_scores_gemma":[0.031279802,0.0007484176,0.0022019898,0.012872694,0.0009878164,0.00965091,0.006222138,0.0017844526,0.002898259],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056622556,0.00068104843,0.013552206,0.0029614714,0.000456474,0.0012689204,0.0057217716,0.004625987,0.031625994,0.074930735,0.05276839,0.8108407],"study_design_scores_gemma":[0.0005974589,0.0007978777,0.02178725,0.0023693242,0.002652776,0.0037820498,0.0070233266,0.17292857,0.12632182,0.09869937,0.5623302,0.0007100297],"about_ca_topic_score_codex":0.009756058,"about_ca_topic_score_gemma":0.011108521,"teacher_disagreement_score":0.013109737,"about_ca_system_score_codex":0.0021313427,"about_ca_system_score_gemma":0.016526526,"threshold_uncertainty_score":0.069331765},"labels":[],"label_agreement":null},{"id":"W4405934475","doi":"10.2196/55277","title":"Creation of Scientific Response Documents for Addressing Product Medical Information Inquiries: Mixed Method Approach Using Artificial Intelligence","year":2024,"lang":"en","type":"article","venue":"JMIR AI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Pfizer (Canada)","funders":"","keywords":"Preprint; Product (mathematics); Data science; Pharmaceutical industry; Computer science; Business; World Wide Web; Medicine","score_opus":0.08223819874697301,"score_gpt":0.4333640955664769,"score_spread":0.3511258968195039,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4405934475","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08052166,0.0020545744,0.8090539,0.0020474293,0.00025750103,0.09260872,0.002617089,0.0018184684,0.009020653],"genre_scores_gemma":[0.083686665,0.00040998362,0.84953916,0.00072143425,0.0000734612,0.06303244,0.00079674105,0.00017731631,0.0015628745],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.7187058,0.24257779,0.015394613,0.010063341,0.012079695,0.0011788092],"domain_scores_gemma":[0.34273317,0.57273257,0.027236553,0.02234714,0.033132076,0.0018185723],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.2037963,0.002600614,0.0024800661,0.013011793,0.0040618656,0.009418691,0.0065068863,0.003295638,0.011553617],"category_scores_gemma":[0.36346903,0.0025207873,0.0035614937,0.009136298,0.002697691,0.006212263,0.008782143,0.0034140314,0.0030232512],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0042920727,0.0046049673,0.020907091,0.018744105,0.0017144034,0.0006420361,0.07038901,0.00584955,0.006216813,0.020977357,0.007059176,0.83860344],"study_design_scores_gemma":[0.010665579,0.0158632,0.06884232,0.017277695,0.0063753286,0.002056511,0.12332664,0.35651654,0.062992774,0.16171768,0.17201495,0.0023508254],"about_ca_topic_score_codex":0.0046310984,"about_ca_topic_score_gemma":0.007775874,"teacher_disagreement_score":0.2037963,"about_ca_system_score_codex":0.0069043287,"about_ca_system_score_gemma":0.012208075,"threshold_uncertainty_score":0.98186094},"labels":[],"label_agreement":null},{"id":"W4406174812","doi":"10.2196/51598","title":"How to Design Electronic Case Report Form (eCRF) Questions to Maximize Semantic Interoperability in Clinical Research","year":2025,"lang":"en","type":"article","venue":"Interactive Journal of Medical Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Interoperability; Computer science; Information retrieval; World Wide Web","score_opus":0.15684430543024733,"score_gpt":0.5448064106503372,"score_spread":0.3879621052200899,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406174812","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0154281845,0.0017688357,0.8946005,0.03736297,0.000821779,0.015093042,0.0017476057,0.004918646,0.028258383],"genre_scores_gemma":[0.031899158,0.00050074124,0.958833,0.0020259453,0.00015962032,0.004094367,0.0008631944,0.00030721386,0.00131682],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.70168805,0.24205193,0.033660606,0.008635767,0.011444493,0.0025190567],"domain_scores_gemma":[0.46309713,0.37331942,0.03076975,0.082269415,0.0447674,0.0057769003],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2871359,0.0015287662,0.0014782485,0.008877294,0.0037204216,0.013480415,0.0054231486,0.0057408214,0.011816639],"category_scores_gemma":[0.43481034,0.0014959088,0.0033632051,0.0057694158,0.008205187,0.025402034,0.011349536,0.0038425948,0.00714192],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060479424,0.0007104422,0.013278771,0.007513679,0.00021565548,0.0008779084,0.037610475,0.0038580147,0.004929112,0.18130955,0.055286955,0.6938046],"study_design_scores_gemma":[0.0005032057,0.0003929297,0.009676166,0.012704417,0.00042719409,0.002222979,0.026845535,0.017537195,0.00980649,0.34335023,0.5760528,0.00048087531],"about_ca_topic_score_codex":0.0029264118,"about_ca_topic_score_gemma":0.0028264984,"teacher_disagreement_score":0.7128641,"about_ca_system_score_codex":0.005203443,"about_ca_system_score_gemma":0.015201385,"threshold_uncertainty_score":0.87908834},"labels":[],"label_agreement":null},{"id":"W4406235660","doi":"10.1016/j.endend.2014.01.021","title":"10.1016/j.endend.2014.01.021","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Property (philosophy); Test (biology); Database; Operating system; Geology","score_opus":0.005943066551645577,"score_gpt":0.19757864433940112,"score_spread":0.19163557778775553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406235660","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008849149,0.008871672,0.033477508,0.010884389,0.0021784208,0.00014775955,0.013231074,0.011123277,0.91123676],"genre_scores_gemma":[0.023105899,0.0036078126,0.022569904,0.0019081851,0.00039426703,0.000108877066,0.010637767,0.0011962968,0.9364709],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994578,0.0000492179,0.00004523066,0.00018213135,0.00017420262,0.00009141551],"domain_scores_gemma":[0.9983491,0.00061507075,0.00018760808,0.00021539137,0.00026943212,0.00036342596],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0014202406,0.0012734582,0.0005515742,0.0026303313,0.0010695894,0.004839672,0.0012527436,0.004102303,0.8487477],"category_scores_gemma":[0.0032656873,0.0005752731,0.0008111781,0.0020598695,0.0014803346,0.00324661,0.0021509123,0.0015438595,0.83852327],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016475267,0.00021375835,0.0065195356,0.00047320814,0.00005725297,0.00039943677,0.00021821771,0.001046805,0.0019142344,0.013358634,0.2595171,0.716117],"study_design_scores_gemma":[0.000048283167,0.00006710758,0.0044320817,0.000894749,0.000057656813,0.001540518,0.0005466803,0.0021065834,0.0013658506,0.01767355,0.97122025,0.00004665612],"about_ca_topic_score_codex":0.0029861063,"about_ca_topic_score_gemma":0.002783976,"teacher_disagreement_score":0.15125233,"about_ca_system_score_codex":0.00088512775,"about_ca_system_score_gemma":0.0013720648,"threshold_uncertainty_score":0.21574312},"labels":[],"label_agreement":null},{"id":"W4406352685","doi":"10.5195/jmla.2025.1936","title":"Algorithmic indexing in MEDLINE frequently overlooks important concepts and may compromise literature search results","year":2025,"lang":"en","type":"article","venue":"Journal of the Medical Library Association JMLA","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Public Health Agency of Canada; Université de Montréal","funders":"McGill University Health Centre; McGill University","keywords":"Search engine indexing; Information retrieval; MEDLINE; Computer science; Medical record; Subject (documents); Index (typography); Data mining; Medicine; Library science; World Wide Web; Radiology","score_opus":0.005347617935268055,"score_gpt":0.2804653613786845,"score_spread":0.27511774344341644,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406352685","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26648402,0.11113527,0.44814882,0.05844226,0.0044828723,0.026315123,0.007979586,0.006766876,0.070245236],"genre_scores_gemma":[0.33539617,0.019944882,0.6176063,0.008974426,0.0015262599,0.009312812,0.0041409093,0.0008279998,0.0022702212],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.63052416,0.19796449,0.10258626,0.0073080016,0.05996378,0.0016533522],"domain_scores_gemma":[0.2798276,0.56133306,0.07424514,0.03345203,0.04978503,0.0013571639],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.27129462,0.0014806879,0.0033671597,0.032546956,0.0034237124,0.015352051,0.0042371606,0.0021164897,0.0039361417],"category_scores_gemma":[0.66835046,0.001301354,0.0026387135,0.032373417,0.005294185,0.012094299,0.007333483,0.0017761328,0.0025849128],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001776926,0.00035885978,0.06064709,0.037886858,0.0016274065,0.00088332087,0.014030273,0.0041029966,0.007607393,0.03577589,0.025471918,0.8098311],"study_design_scores_gemma":[0.002001764,0.0038510805,0.17291433,0.08826359,0.005190854,0.01116245,0.026058286,0.036724664,0.02262796,0.21618688,0.4133962,0.0016219732],"about_ca_topic_score_codex":0.0032930202,"about_ca_topic_score_gemma":0.0054088724,"teacher_disagreement_score":0.9846479,"about_ca_system_score_codex":0.0060773664,"about_ca_system_score_gemma":0.016943125,"threshold_uncertainty_score":0.89862347},"labels":[],"label_agreement":null},{"id":"W4406352735","doi":"10.5195/jmla.2025.1972","title":"Filtering failure: the impact of automated indexing in Medline on retrieval of human studies for knowledge synthesis","year":2025,"lang":"en","type":"article","venue":"Journal of the Medical Library Association JMLA","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Search engine indexing; Information retrieval; MEDLINE; Computer science; Data science; Biology","score_opus":0.021491924303960182,"score_gpt":0.36333240583752263,"score_spread":0.34184048153356245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406352735","genre_codex":"review","genre_gemma":"empirical","domain_codex":"methods","domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13926212,0.30972025,0.30428478,0.11085466,0.014641763,0.04195473,0.028407108,0.008496885,0.04237769],"genre_scores_gemma":[0.5706272,0.03180327,0.3205366,0.029508589,0.0034476242,0.03054082,0.007985202,0.0023160097,0.00323469],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.17789692,0.47629216,0.25448883,0.016662503,0.07157945,0.0030801904],"domain_scores_gemma":[0.023251818,0.8898751,0.04290024,0.022657461,0.020455066,0.0008601984],"candidate_categories":["metaresearch","scholarly_communication"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.67402303,0.0041090413,0.009605333,0.03692408,0.0042671016,0.013314212,0.007305851,0.00751089,0.010862211],"category_scores_gemma":[0.90669835,0.0028410289,0.012010196,0.033524387,0.006415856,0.015179358,0.008881738,0.0036201335,0.0022223133],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013921383,0.00036044503,0.06460665,0.26207438,0.026682312,0.0017369129,0.021262752,0.00437331,0.0065924674,0.014828785,0.07887541,0.50468516],"study_design_scores_gemma":[0.009456526,0.0052461354,0.14512068,0.37988898,0.07009168,0.006726619,0.0070734913,0.0382031,0.019163972,0.061690148,0.25473687,0.0026018557],"about_ca_topic_score_codex":0.012724517,"about_ca_topic_score_gemma":0.013795358,"teacher_disagreement_score":0.9866858,"about_ca_system_score_codex":0.012768031,"about_ca_system_score_gemma":0.024959724,"threshold_uncertainty_score":0.40198767},"labels":[],"label_agreement":null},{"id":"W4406441607","doi":"10.1016/s0029-7437(06)71683-4","title":"10.1016/s0029-7437(06)71683-4","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Materials science","score_opus":0.00720098715786286,"score_gpt":0.20384627860626187,"score_spread":0.196645291448399,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406441607","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00049224315,0.0005245119,0.0018597848,0.0005345171,0.00043041605,0.00013702088,0.0016079914,0.0019023401,0.9925113],"genre_scores_gemma":[0.00058879744,0.00021284542,0.0008105831,0.0001988,0.00008209984,0.00007593682,0.0008439956,0.00031652264,0.99687046],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990182,0.00007754618,0.00008974315,0.00034355512,0.00027507983,0.00019583489],"domain_scores_gemma":[0.9962528,0.0011149134,0.00022748971,0.00052864273,0.00065922324,0.001217014],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0018548267,0.0034275022,0.0024061392,0.0035633962,0.002697337,0.0055906167,0.0044750837,0.005937423,0.98998904],"category_scores_gemma":[0.0025239752,0.001354748,0.0021376936,0.004310666,0.002891016,0.0072576366,0.00429555,0.0037581262,0.9940546],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030551752,0.00022809152,0.0008296215,0.0006373535,0.000056811354,0.00024170017,0.00012212597,0.00057617197,0.002540886,0.006651051,0.32962766,0.658183],"study_design_scores_gemma":[0.00006274528,0.00011827255,0.0007524413,0.0003464684,0.000018695919,0.0002871033,0.00012938489,0.0004080428,0.00051240035,0.000939127,0.99639004,0.000035156423],"about_ca_topic_score_codex":0.004718701,"about_ca_topic_score_gemma":0.003986344,"teacher_disagreement_score":0.010010958,"about_ca_system_score_codex":0.0014942845,"about_ca_system_score_gemma":0.0017579241,"threshold_uncertainty_score":0.014279306},"labels":[],"label_agreement":null},{"id":"W4406472716","doi":"10.1016/s1558-0164(06)70099-2","title":"10.1016/s1558-0164(06)70099-2","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine","score_opus":0.005929070311962406,"score_gpt":0.19456264788432387,"score_spread":0.18863357757236146,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406472716","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004032928,0.00036445446,0.0016120692,0.0005104924,0.00036036276,0.00011423828,0.0012215528,0.0018786766,0.9935348],"genre_scores_gemma":[0.0005080307,0.00014937778,0.00061290467,0.00020750171,0.00006934924,0.000050623017,0.00064243283,0.00025087484,0.9975089],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991487,0.00006553726,0.00007354048,0.0003087083,0.00023804177,0.00016542317],"domain_scores_gemma":[0.99710566,0.0006975375,0.00017174827,0.00044562598,0.0005648117,0.0010145608],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0014968386,0.0026205797,0.001654981,0.0030062646,0.0024707655,0.004669788,0.0032877987,0.005290841,0.9901631],"category_scores_gemma":[0.002327716,0.0010160194,0.0016293313,0.0030785627,0.0019066826,0.0062352996,0.004087165,0.002662962,0.993979],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028265858,0.00019667183,0.0008253663,0.00041408406,0.00003896099,0.00024218077,0.00010344424,0.00044903113,0.0020917584,0.005752057,0.3607912,0.6288127],"study_design_scores_gemma":[0.000048658505,0.00008564497,0.0005643765,0.00022074368,0.000014177567,0.00028676112,0.000099096884,0.00027453373,0.00040066245,0.0008050207,0.99717665,0.000023655737],"about_ca_topic_score_codex":0.0035282853,"about_ca_topic_score_gemma":0.0032913194,"teacher_disagreement_score":0.009836912,"about_ca_system_score_codex":0.0011849751,"about_ca_system_score_gemma":0.0013005239,"threshold_uncertainty_score":0.014031053},"labels":[],"label_agreement":null},{"id":"W4406518288","doi":"10.1016/s0029-7437(08)70292-1","title":"10.1016/s0029-7437(08)70292-1","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Osteoporosis; Medicine; Psychology; Computer science; Internal medicine","score_opus":0.0071642813954135055,"score_gpt":0.20361468956000392,"score_spread":0.19645040816459042,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406518288","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004560585,0.00042050314,0.0017026619,0.00052736676,0.00038469152,0.00012894181,0.0013281126,0.0017552718,0.9932962],"genre_scores_gemma":[0.00055249094,0.00019585692,0.00079437764,0.00019786197,0.00007771905,0.00007525766,0.00075184525,0.0003211145,0.9970336],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99911016,0.0000686631,0.000077552875,0.00031461078,0.00024639908,0.00018257432],"domain_scores_gemma":[0.9965023,0.0010185908,0.00020084831,0.00047790492,0.00059604814,0.0012042283],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0016725097,0.0031379517,0.0021735686,0.0034027225,0.0025296973,0.005014629,0.0041098883,0.005714172,0.9895446],"category_scores_gemma":[0.0025351522,0.0011938352,0.0019230647,0.0038064735,0.0025506152,0.006852047,0.0040523456,0.0033759996,0.9936132],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028358013,0.00022531094,0.00082838116,0.00053586875,0.00004534405,0.0002393819,0.0001233126,0.00053875265,0.0023547115,0.006427249,0.34277245,0.6456256],"study_design_scores_gemma":[0.00005570373,0.00010849206,0.0007479479,0.00031340597,0.000016043836,0.00026419258,0.00012666563,0.00038981024,0.0004381916,0.0008904768,0.99661946,0.000029518675],"about_ca_topic_score_codex":0.0045011737,"about_ca_topic_score_gemma":0.0036372754,"teacher_disagreement_score":0.01045543,"about_ca_system_score_codex":0.0013843417,"about_ca_system_score_gemma":0.0016879098,"threshold_uncertainty_score":0.01491338},"labels":[],"label_agreement":null},{"id":"W4406546290","doi":"10.1016/s1544-8800(05)70203-7","title":"10.1016/s1544-8800(05)70203-7","year":2000,"lang":"en","type":"article","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Field (mathematics); Medicine; Mathematics","score_opus":0.006167768870811256,"score_gpt":0.1955757150397952,"score_spread":0.18940794616898396,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406546290","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00046546536,0.00045824968,0.0015255816,0.0005260592,0.0004013213,0.0001310741,0.0014886407,0.0021292956,0.99287415],"genre_scores_gemma":[0.0005237637,0.00019938013,0.0006521301,0.0002502519,0.000082878294,0.00006882249,0.00075558433,0.00028451823,0.9971827],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991159,0.00006381547,0.000079403275,0.00031985057,0.00024292442,0.0001780763],"domain_scores_gemma":[0.9963911,0.0009839071,0.00020448533,0.000508409,0.00068542984,0.0012265131],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0016205458,0.0031381007,0.0020534045,0.0032276362,0.0026219934,0.0047239773,0.0036722405,0.0059192367,0.9900572],"category_scores_gemma":[0.0024066004,0.0011072932,0.0018834377,0.0034712828,0.0019681416,0.0068826135,0.0043472215,0.0030880547,0.99403375],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033100852,0.00021092479,0.00076722674,0.00047716897,0.000040946878,0.00020810831,0.00009718324,0.00045639923,0.0021669308,0.0047147083,0.35498205,0.6355474],"study_design_scores_gemma":[0.00007114243,0.0001251527,0.0007746462,0.00029019284,0.000018511539,0.00027978883,0.000113086004,0.00034366356,0.00047140813,0.00076714956,0.9967128,0.000032418196],"about_ca_topic_score_codex":0.0045049484,"about_ca_topic_score_gemma":0.0038528682,"teacher_disagreement_score":0.00994283,"about_ca_system_score_codex":0.0012953855,"about_ca_system_score_gemma":0.0014410946,"threshold_uncertainty_score":0.01418227},"labels":[],"label_agreement":null},{"id":"W4406641470","doi":"10.1016/b978-2-294-71511-2.00006-4","title":"10.1016/b978-2-294-71511-2.00006-4","year":2000,"lang":"en","type":"book-chapter","venue":"Time to knit","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"BI-RADS; Medicine; Mammography; Cancer; Breast cancer; Internal medicine","score_opus":0.009073331753210791,"score_gpt":0.1942506842305501,"score_spread":0.18517735247733932,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406641470","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002697145,0.00063228275,0.0024897233,0.0005397208,0.00024297851,0.00005716097,0.0014073096,0.002210612,0.99215066],"genre_scores_gemma":[0.0007243184,0.00033109522,0.0009620794,0.00017978121,0.000047093832,0.0000467269,0.00084274786,0.00045032342,0.99641585],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994475,0.000034286546,0.000041218,0.00018654621,0.00019419688,0.000096169366],"domain_scores_gemma":[0.998106,0.0006504369,0.0001302421,0.00029566637,0.0003052174,0.00051236677],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0011070358,0.0020036583,0.0014293691,0.0019228158,0.0012010611,0.005061876,0.0027451746,0.004016984,0.9822872],"category_scores_gemma":[0.0020476796,0.0008308158,0.001061539,0.002154973,0.0012086527,0.0052297916,0.0034508873,0.00228476,0.99179196],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013573979,0.00010886266,0.0003976805,0.00038470683,0.00002006411,0.00011257442,0.00007241402,0.0003899957,0.0021200546,0.006119156,0.31899965,0.6711391],"study_design_scores_gemma":[0.000019790114,0.00003948607,0.00051220995,0.00026863103,0.000010217494,0.00026421066,0.00006943775,0.00016930782,0.00041058334,0.0015540196,0.9966659,0.00001626162],"about_ca_topic_score_codex":0.0021246215,"about_ca_topic_score_gemma":0.002018522,"teacher_disagreement_score":0.017712772,"about_ca_system_score_codex":0.0009780369,"about_ca_system_score_gemma":0.00073834165,"threshold_uncertainty_score":0.025265098},"labels":[],"label_agreement":null},{"id":"W4406661144","doi":"10.1016/j.reprotox.2025.108838","title":"Comprehensive digital documentation of teratology studies","year":2025,"lang":"en","type":"article","venue":"Reproductive Toxicology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Nexen (Canada)","funders":"","keywords":"Teratology; Documentation; Medicine; Computer science; Biology; Pregnancy","score_opus":0.02832477946680387,"score_gpt":0.362573629670726,"score_spread":0.3342488502039221,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406661144","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019185996,0.017756114,0.07961177,0.0034601863,0.00056166,0.00062104786,0.82903236,0.014352199,0.0354186],"genre_scores_gemma":[0.07589908,0.025420044,0.11994842,0.0014684369,0.0003563282,0.0006537296,0.76111764,0.001894033,0.013242168],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983687,0.00028787213,0.0005937295,0.00020749803,0.00048681357,0.000055478296],"domain_scores_gemma":[0.9815248,0.010695565,0.0022918284,0.0022147822,0.0027672637,0.0005057289],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0024631675,0.0009614396,0.0007252286,0.014556386,0.0005636087,0.0028435048,0.0009800306,0.0012166434,0.017449817],"category_scores_gemma":[0.016022839,0.00042911476,0.00059144787,0.011206459,0.0004099016,0.0027655133,0.0017091378,0.001001742,0.0069368607],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047307683,0.00024491697,0.012763005,0.017551929,0.00024962967,0.0024822997,0.001861549,0.00609605,0.012360129,0.022494273,0.34421858,0.5792046],"study_design_scores_gemma":[0.00006108099,0.000039073126,0.010100187,0.0029587557,0.00027837022,0.0020433988,0.0004764668,0.003848491,0.006632942,0.014244457,0.95925343,0.000063401196],"about_ca_topic_score_codex":0.004216673,"about_ca_topic_score_gemma":0.0072783446,"teacher_disagreement_score":0.99753684,"about_ca_system_score_codex":0.000799957,"about_ca_system_score_gemma":0.0038758933,"threshold_uncertainty_score":0.058375478},"labels":[],"label_agreement":null},{"id":"W4406822659","doi":"10.2196/68704","title":"Improving Phenotyping of Patients With Immune-Mediated Inflammatory Diseases Through Automated Processing of Discharge Summaries: Multicenter Cohort Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Medicine; Preprint; Cohort; Cohort study; Intensive care medicine; Immune system; Immunology; Computer science; Internal medicine; World Wide Web","score_opus":0.004565167902020977,"score_gpt":0.2652871187780031,"score_spread":0.2607219508759821,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406822659","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9980856,0.000056176585,0.0005695699,0.00004264085,0.0000073721308,0.000079579244,0.0010793029,0.0000139008735,0.00006588667],"genre_scores_gemma":[0.9951278,0.000073154886,0.001460484,0.000055799723,0.000024093802,0.0001050501,0.003055407,0.000012724496,0.000085403015],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9977239,0.0009118667,0.00023403716,0.0007286298,0.00021920771,0.00018236326],"domain_scores_gemma":[0.9926795,0.001947104,0.0020760691,0.0015658608,0.001132901,0.0005984716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004848144,0.00052023627,0.00053805095,0.001113453,0.0005059294,0.0010730675,0.0005639338,0.0005912706,0.00077693275],"category_scores_gemma":[0.012410804,0.00036395993,0.0011170512,0.0009198227,0.00032684996,0.0008518342,0.0009695708,0.0007813001,0.0002728613],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011155602,0.000405741,0.9927314,0.00003119968,0.00028183058,0.00012641717,0.0003164216,0.00020659655,0.00066684134,0.000035081357,0.00048141027,0.0036015017],"study_design_scores_gemma":[0.00018678933,0.0013212452,0.9922315,0.000026267753,0.00031186515,0.00036368167,0.00065772363,0.0033251673,0.0006016015,0.00007611707,0.00086878415,0.000029184206],"about_ca_topic_score_codex":0.004131162,"about_ca_topic_score_gemma":0.0038972376,"teacher_disagreement_score":0.004848144,"about_ca_system_score_codex":0.0005574007,"about_ca_system_score_gemma":0.0008389589,"threshold_uncertainty_score":0.025639713},"labels":[],"label_agreement":null},{"id":"W4406916996","doi":"10.1016/j.matbio.2025.01.007","title":"George R. Martin: Pioneering matrix biologist","year":2025,"lang":"en","type":"article","venue":"Matrix Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Saskatchewan","funders":"Canadian Institutes of Health Research; National Institute for Dementia Research; National Institutes of Health","keywords":"George (robot); Biologist; Philosophy; Matrix (chemical analysis); Environmental ethics; Art history; Art; Chemistry; Biology; Genetics","score_opus":0.008937574607784812,"score_gpt":0.3264261930731696,"score_spread":0.3174886184653848,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4406916996","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008230955,0.10875525,0.4528475,0.33236653,0.036714748,0.0003720326,0.0034030394,0.010404334,0.046905685],"genre_scores_gemma":[0.060397726,0.09477632,0.46772292,0.056856275,0.019366916,0.0005820702,0.00401684,0.003729197,0.2925518],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968472,0.0005633113,0.0001436836,0.0008487741,0.0014154885,0.00018147122],"domain_scores_gemma":[0.98779804,0.004298703,0.00046215567,0.0008978396,0.0036812017,0.0028620178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0073080203,0.0013548402,0.0010706567,0.004421553,0.0012311565,0.003192914,0.001646226,0.002732783,0.016639078],"category_scores_gemma":[0.012862192,0.0009196421,0.00093095226,0.0021871596,0.0022023858,0.00443446,0.0019945675,0.0059106993,0.016859325],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00035425922,0.00019247968,0.0021191684,0.00083611073,0.0001298538,0.0005653442,0.0006845187,0.00080827816,0.0220988,0.045931213,0.6027004,0.32357964],"study_design_scores_gemma":[0.000040403156,0.00010551043,0.0010401972,0.00020481585,0.000038390062,0.0011681252,0.00017785004,0.0016201795,0.01045217,0.030310731,0.9547705,0.000071038696],"about_ca_topic_score_codex":0.002571696,"about_ca_topic_score_gemma":0.0038081594,"teacher_disagreement_score":0.016639078,"about_ca_system_score_codex":0.001311799,"about_ca_system_score_gemma":0.0031756114,"threshold_uncertainty_score":0.055663288},"labels":[],"label_agreement":null},{"id":"W4407006479","doi":"10.1093/database/baae097","title":"Helping authors produce FAIR taxonomic data: evaluation of an author-driven phenotype data production prototype","year":2025,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Museum of Nature; Agriculture and Agri-Food Canada; Université de Montréal; Ministry of Natural Resources and Forestry; University of Manitoba; University of Ottawa","funders":"National Science Foundation","keywords":"Computer science; Ontology; Workflow; Character (mathematics); Vocabulary; World Wide Web; Information retrieval; Data science; Artificial intelligence; Database","score_opus":0.15365448967522466,"score_gpt":0.40095387990632747,"score_spread":0.2472993902311028,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407006479","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6097322,0.00021622158,0.3185634,0.00080378365,0.00027331678,0.00287163,0.0019594603,0.05621293,0.009367017],"genre_scores_gemma":[0.35090792,0.00020327259,0.6250633,0.00041686997,0.00004119101,0.0015337478,0.004665364,0.00505333,0.012115028],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9950723,0.0021764534,0.00053512363,0.00079073355,0.00125074,0.00017466057],"domain_scores_gemma":[0.9543659,0.032527972,0.0006861633,0.006165725,0.004767669,0.0014866733],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.011042111,0.0009924959,0.0008557152,0.0009416524,0.0007178026,0.002930355,0.0044433502,0.0023104313,0.0067629353],"category_scores_gemma":[0.03745289,0.00065096724,0.00081152597,0.00067550415,0.0009467748,0.0037316794,0.002837311,0.001256116,0.0027161958],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0068076174,0.008674382,0.021788608,0.0032509442,0.00029175018,0.0077325907,0.052858334,0.022331633,0.20055524,0.009319038,0.042217202,0.62417275],"study_design_scores_gemma":[0.0035701953,0.013212218,0.021729646,0.00083318754,0.00064173917,0.0062659206,0.015342622,0.43678236,0.23401554,0.010531795,0.25617552,0.00089922966],"about_ca_topic_score_codex":0.0016179911,"about_ca_topic_score_gemma":0.0016568556,"teacher_disagreement_score":0.99555665,"about_ca_system_score_codex":0.0010215426,"about_ca_system_score_gemma":0.0013569716,"threshold_uncertainty_score":0.058396935},"labels":[],"label_agreement":null},{"id":"W4407252450","doi":"10.1038/s41698-025-00824-w","title":"An automatic pipeline for temporal monitoring of radiotherapy-induced toxicities in head and neck cancer patients","year":2025,"lang":"en","type":"article","venue":"npj Precision Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Jewish General Hospital; McGill University; McGill University Health Centre","funders":"Canada Research Chairs","keywords":"Radiation therapy; Head and neck cancer; Pipeline (software); Medicine; Data extraction; Toxicity; Cancer; Computer science; Internal medicine; MEDLINE","score_opus":0.027052440742807667,"score_gpt":0.3844105452528156,"score_spread":0.35735810451000793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407252450","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15732974,0.004046456,0.6348697,0.0030590605,0.0005587638,0.0014909452,0.07751029,0.1151276,0.0060074236],"genre_scores_gemma":[0.33917844,0.001265494,0.546219,0.0007613168,0.00024124175,0.00093420886,0.1047239,0.001342187,0.0053342558],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999175,0.00010301588,0.00012877809,0.00035861714,0.00014669784,0.00008783327],"domain_scores_gemma":[0.9983741,0.0006506792,0.00028278425,0.00018648096,0.00039186567,0.00011405681],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013138896,0.0012115482,0.0007031196,0.0033244116,0.00060695,0.0011550286,0.00094212976,0.0009972122,0.0035377769],"category_scores_gemma":[0.0039750044,0.00042059325,0.0013662545,0.0018894891,0.00025333842,0.0011294138,0.0012950932,0.00088419515,0.003095918],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012228257,0.0004950487,0.03822929,0.0012149331,0.0002457681,0.0021123234,0.0007931887,0.011004293,0.07236919,0.0023479608,0.07335853,0.79660666],"study_design_scores_gemma":[0.0003836112,0.0006091644,0.06547945,0.00034050434,0.00061716704,0.004705337,0.001045383,0.63026494,0.1412329,0.018872177,0.13620387,0.0002456101],"about_ca_topic_score_codex":0.007725739,"about_ca_topic_score_gemma":0.011097601,"teacher_disagreement_score":0.007725739,"about_ca_system_score_codex":0.0010291233,"about_ca_system_score_gemma":0.0024369296,"threshold_uncertainty_score":0.0153615475},"labels":[],"label_agreement":null},{"id":"W4407253515","doi":"10.7554/elife.105565.1.sa3","title":"eLife Assessment: Interpretable Protein-DNA Interactions Captured by Structure-Sequence Optimization","year":2025,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Sequence (biology); Computational biology; DNA; Computer science; DNA sequencing; Artificial intelligence; Biology; Genetics","score_opus":0.017654183877419308,"score_gpt":0.3434417609449189,"score_spread":0.3257875770674996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407253515","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.07093781,0.0018371321,0.89396983,0.002385269,0.00033206103,0.0002511136,0.003435036,0.010375097,0.016476618],"genre_scores_gemma":[0.46920115,0.0015713342,0.5101602,0.0007421686,0.00013320644,0.00062818156,0.005884826,0.002534948,0.009143946],"study_design_codex":"simulation_or_modeling","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981186,0.00061247265,0.000088942616,0.00027159706,0.00079760543,0.0001108098],"domain_scores_gemma":[0.99678326,0.0015496833,0.0003110119,0.00049132446,0.0007382786,0.00012644262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034331048,0.00083718775,0.001239483,0.0014094764,0.0005953012,0.0018828963,0.0029943674,0.0014898443,0.006915654],"category_scores_gemma":[0.013003038,0.0004989331,0.0007368469,0.0012171431,0.00097329385,0.0019204225,0.0015765701,0.0014060563,0.0030486088],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041468293,0.00015460239,0.009117502,0.0010488112,0.0002956118,0.00051550666,0.00040848053,0.63093185,0.023009036,0.07444099,0.037142616,0.22252032],"study_design_scores_gemma":[0.00005897073,0.000049133705,0.0010351714,0.00007474866,0.00004319434,0.00010392809,0.00010710377,0.8879623,0.016381018,0.06286372,0.03128031,0.000040550134],"about_ca_topic_score_codex":0.0026756795,"about_ca_topic_score_gemma":0.004194569,"teacher_disagreement_score":0.006915654,"about_ca_system_score_codex":0.0011952189,"about_ca_system_score_gemma":0.0028688102,"threshold_uncertainty_score":0.023135185},"labels":[],"label_agreement":null},{"id":"W4407639229","doi":"10.1109/bip63158.2024.10885387","title":"Relation Extraction from Unstructured Species Descriptions Using TaxonNERD and Llama 7B","year":2024,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Ministerio de Ciencia Tecnología y Telecomunicaciones; Instituto Tecnológico de Costa Rica; Consejo Superior Universitario Centroamericano; International Development Research Centre","keywords":"Relation (database); Computer science; Extraction (chemistry); Relationship extraction; Data mining; Chromatography; Chemistry","score_opus":0.03603538874174903,"score_gpt":0.2889795448439157,"score_spread":0.25294415610216664,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407639229","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08346624,0.0049127517,0.5509178,0.0022921786,0.00036983972,0.0021315846,0.2715535,0.07135947,0.012996668],"genre_scores_gemma":[0.09102784,0.0011336444,0.67658836,0.0003552574,0.00003955751,0.0011560052,0.22638837,0.0009015947,0.0024094412],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99845064,0.00032460678,0.0003873577,0.00039513136,0.00037036487,0.00007189635],"domain_scores_gemma":[0.9965347,0.001905429,0.0003846142,0.00059736095,0.00046923725,0.00010859742],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019985333,0.001312792,0.0008837533,0.010611161,0.001175443,0.002430995,0.0011043987,0.0012120121,0.004892256],"category_scores_gemma":[0.009183114,0.00064437214,0.0020459,0.0058386875,0.00050670886,0.003508878,0.0025395693,0.0010796944,0.0028805672],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008786928,0.00061226217,0.056795146,0.008162904,0.00088742585,0.003250615,0.0043632705,0.027052896,0.035031844,0.03976779,0.1419825,0.68121463],"study_design_scores_gemma":[0.00020614648,0.00016732122,0.029905047,0.0011599351,0.0003705615,0.0024666928,0.002178954,0.3318225,0.029352799,0.04031582,0.5618592,0.00019497819],"about_ca_topic_score_codex":0.0150329685,"about_ca_topic_score_gemma":0.034395244,"teacher_disagreement_score":0.0150329685,"about_ca_system_score_codex":0.0015355451,"about_ca_system_score_gemma":0.0025114731,"threshold_uncertainty_score":0.029890954},"labels":[],"label_agreement":null},{"id":"W4407715564","doi":"10.1503/cmaj.241351-f","title":"Indicateurs pour l’évaluation de l’interopérabilité des données sur la santé au Canada","year":2025,"lang":"fr","type":"article","venue":"Canadian Medical Association Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"Ottawa Public Health","funders":"","keywords":"Medicine; Computer science","score_opus":0.010424071782888187,"score_gpt":0.2592751325764931,"score_spread":0.24885106079360492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407715564","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4977231,0.028471474,0.15679574,0.06976688,0.0016198346,0.004353072,0.12823415,0.008504673,0.10453113],"genre_scores_gemma":[0.7155267,0.006332797,0.20756513,0.0022637227,0.00019208308,0.0016380309,0.05752583,0.0010310949,0.007924683],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9084422,0.027882086,0.007541757,0.0063219597,0.046476636,0.003335368],"domain_scores_gemma":[0.7261023,0.12748332,0.008820694,0.010771605,0.12178213,0.0050400374],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.08031088,0.002325813,0.0018197178,0.0140994685,0.005305726,0.018857278,0.003087789,0.0022325867,0.006566215],"category_scores_gemma":[0.25123248,0.0009288423,0.0033435377,0.019652193,0.0024737264,0.007578823,0.0062881466,0.0044135605,0.0014951014],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029282619,0.00092296104,0.37158555,0.0069897366,0.004015769,0.0006422623,0.016999653,0.039979782,0.012556356,0.051259328,0.081342354,0.41077793],"study_design_scores_gemma":[0.0006484119,0.0011844527,0.4576581,0.010929328,0.0032366456,0.0005996818,0.037368026,0.23169622,0.027849726,0.028376771,0.19944894,0.0010036831],"about_ca_topic_score_codex":0.9006585,"about_ca_topic_score_gemma":0.8438519,"teacher_disagreement_score":0.94905216,"about_ca_system_score_codex":0.050947864,"about_ca_system_score_gemma":0.07755037,"threshold_uncertainty_score":0.4247296},"labels":[],"label_agreement":null},{"id":"W4407753198","doi":"10.3233/shti250022","title":"Bridging the Evidence-to-Practice Gap at Scale: Evolving Evidence2Practice Ontario by Using the Pan-Canadian HALO Framework","year":2025,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto; Canada Health Infoway","funders":"","keywords":"Bridging (networking); Interoperability; Scalability; Halo; Protocol (science); Clinical decision support system; Clinical Practice; Computer science; Scale (ratio); Decision support system; Process management; Data science; Knowledge management; Business; Medicine; Nursing; World Wide Web; Computer security; Database; Data mining; Geography","score_opus":0.06415442099447202,"score_gpt":0.39857900860118434,"score_spread":0.33442458760671234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407753198","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.042882476,0.02251762,0.11564274,0.6132504,0.0024909638,0.0036055834,0.011301122,0.003347722,0.18496144],"genre_scores_gemma":[0.45861122,0.017743839,0.43756315,0.048959967,0.0009087184,0.0038763,0.0110123325,0.0017746065,0.019549802],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.8510958,0.050905306,0.015238734,0.0063705225,0.06595072,0.010438961],"domain_scores_gemma":[0.46819523,0.19271429,0.018715978,0.034000706,0.24749273,0.03888111],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.18092601,0.0010908045,0.0015775092,0.016608832,0.010086439,0.021557707,0.0088126585,0.0041828696,0.010431059],"category_scores_gemma":[0.31210923,0.0016627223,0.0022390403,0.019512145,0.009667564,0.014582414,0.027270116,0.005942142,0.0018448655],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006041821,0.00025544353,0.04435533,0.008241558,0.0006441991,0.000594281,0.01447687,0.0053666853,0.0026421016,0.2856854,0.20981233,0.42732167],"study_design_scores_gemma":[0.0002623328,0.00024608115,0.044651344,0.012103476,0.0005865241,0.00030038733,0.011196453,0.010436292,0.0026166341,0.09177005,0.8253646,0.0004657598],"about_ca_topic_score_codex":0.8918635,"about_ca_topic_score_gemma":0.9391626,"teacher_disagreement_score":0.8724515,"about_ca_system_score_codex":0.1275485,"about_ca_system_score_gemma":0.43640217,"threshold_uncertainty_score":0.9568396},"labels":[],"label_agreement":null},{"id":"W4407753199","doi":"10.3233/shti250019","title":"A Guide for Implementing an A.I-Driven Initiative in Rural Northern Ontario","year":2025,"lang":"en","type":"article","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Nova Scotia Health Authority; University of Toronto","funders":"","keywords":"Pulmonary embolism; Clinical decision support system; Medicine; Medical emergency; Rural area; Electronic health record; Rural health; Environmental planning; Geography; Computer science; Health care; Decision support system; Political science; Pathology; Surgery; Artificial intelligence","score_opus":0.047085340559278516,"score_gpt":0.3922306961807685,"score_spread":0.34514535562148996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407753199","genre_codex":"other","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026416184,0.003126476,0.17025003,0.11277786,0.0021697283,0.019345725,0.009645324,0.019728424,0.63654023],"genre_scores_gemma":[0.054734584,0.0043553454,0.45461878,0.013252382,0.00025885212,0.008864849,0.0031334166,0.0020785215,0.45870328],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970728,0.0006622761,0.00021684518,0.00019320998,0.0012478958,0.0006069898],"domain_scores_gemma":[0.9901721,0.0006409314,0.00043634803,0.0003906807,0.0047216257,0.0036382922],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006072049,0.00073835114,0.00042572507,0.0017496579,0.006348865,0.004457084,0.0027863916,0.0023164484,0.06963068],"category_scores_gemma":[0.0075975056,0.00088572304,0.0006509927,0.0023217604,0.0019291992,0.0020222794,0.0030681111,0.0015048913,0.03595394],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008926615,0.00037987166,0.0071922583,0.0004810028,0.000008982665,0.0014185667,0.007942896,0.0007989604,0.005426833,0.007411074,0.70208555,0.2667648],"study_design_scores_gemma":[0.000059936614,0.0001607247,0.010593723,0.00025870098,0.000008129839,0.00024499948,0.0053896885,0.0006384322,0.0006126436,0.0015107583,0.98046976,0.000052482654],"about_ca_topic_score_codex":0.6281042,"about_ca_topic_score_gemma":0.8641855,"teacher_disagreement_score":0.3718958,"about_ca_system_score_codex":0.020164818,"about_ca_system_score_gemma":0.07506331,"threshold_uncertainty_score":0.7481719},"labels":[],"label_agreement":null},{"id":"W4407828098","doi":"10.1212/nxg.0000000000200246","title":"The Neurodegenerative Disease Knowledge Portal","year":2025,"lang":"en","type":"review","venue":"Neurology Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"National Human Genome Research Institute","keywords":"Disease; Medicine; Pathology","score_opus":0.02836067930000892,"score_gpt":0.34824866040668,"score_spread":0.3198879811066711,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407828098","genre_codex":"dataset","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00648449,0.0068308976,0.11857994,0.010886468,0.0010434954,0.0010639107,0.6469012,0.10886596,0.09934367],"genre_scores_gemma":[0.029580193,0.0070481943,0.11372582,0.005302983,0.0005552763,0.00097233086,0.81695145,0.007956477,0.01790719],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99800915,0.0004684938,0.00039267057,0.0003473344,0.0006199039,0.00016234949],"domain_scores_gemma":[0.99254614,0.002769998,0.0005107365,0.0018361611,0.0010814073,0.0012554179],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004321755,0.0011711328,0.0014199094,0.0065823407,0.000976563,0.007121776,0.0030358082,0.002722572,0.052547738],"category_scores_gemma":[0.014346466,0.0006736953,0.0008823591,0.007952773,0.0005324051,0.0057440777,0.008243264,0.0025642128,0.0446976],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005534268,0.00017735793,0.003658186,0.0024387476,0.00012921117,0.001541604,0.0005973946,0.0012576322,0.002151671,0.026695164,0.77524686,0.18555275],"study_design_scores_gemma":[0.00015746144,0.00003924324,0.0018435982,0.0005082889,0.000036646183,0.00080510735,0.00019391802,0.0017771252,0.0011753689,0.030414667,0.9629899,0.00005857149],"about_ca_topic_score_codex":0.0031309952,"about_ca_topic_score_gemma":0.004412463,"teacher_disagreement_score":0.052547738,"about_ca_system_score_codex":0.0012898939,"about_ca_system_score_gemma":0.0036167726,"threshold_uncertainty_score":0.17578971},"labels":[],"label_agreement":null},{"id":"W4407834748","doi":"10.3233/978-1-58603-979-0-414","title":"eDoc Evaluation &amp;ndash; At Eighteen Months into the Challenge","year":2009,"lang":"en","type":"book-chapter","venue":"Studies in health technology and informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.06739027356147141,"score_gpt":0.37198401168598794,"score_spread":0.3045937381245165,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4407834748","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25233516,0.0065683518,0.069206625,0.108348474,0.017776962,0.08005856,0.02566514,0.005602147,0.43443856],"genre_scores_gemma":[0.3922681,0.0016864905,0.11804411,0.020040076,0.0013905793,0.033149984,0.017126387,0.0017770047,0.41451737],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9698019,0.012438757,0.0017461343,0.0014785836,0.01104735,0.0034871818],"domain_scores_gemma":[0.8940733,0.009144728,0.0021648428,0.007969604,0.07269802,0.013949485],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05487513,0.0006994855,0.00071374094,0.0025010435,0.0039050314,0.008983505,0.0023605453,0.0019385603,0.022521773],"category_scores_gemma":[0.06441611,0.00049238134,0.0007844888,0.0017183871,0.0013630755,0.0036932267,0.0057397713,0.0026894673,0.008797373],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001588624,0.0038214484,0.008223603,0.00085798325,0.00006300569,0.00026148366,0.0052785473,0.0007849838,0.0020248292,0.008268451,0.3618696,0.6069575],"study_design_scores_gemma":[0.0008097446,0.0044247108,0.031885482,0.0012436128,0.00006804166,0.00019335374,0.010013206,0.0023814612,0.0067627733,0.0032620046,0.9388148,0.00014073761],"about_ca_topic_score_codex":0.02288332,"about_ca_topic_score_gemma":0.05531714,"teacher_disagreement_score":0.05487513,"about_ca_system_score_codex":0.0106372945,"about_ca_system_score_gemma":0.019320054,"threshold_uncertainty_score":0.2902109},"labels":[],"label_agreement":null},{"id":"W4408039465","doi":"10.62617/mcb1326","title":"Weakly-supervised natural language processing with BERT-Clinical for automated lesion information extraction from free-text MRI reports in multiple sclerosis patients","year":2025,"lang":"en","type":"article","venue":"Molecular & cellular biomechanics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Information extraction; Computer science; Text messaging; Natural language processing; Artificial intelligence; Multiple sclerosis; Lesion; Medicine; Information retrieval; Pattern recognition (psychology); Pathology; World Wide Web","score_opus":0.015654928685932906,"score_gpt":0.27368052551820765,"score_spread":0.2580255968322748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408039465","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6195887,0.0029039034,0.35350356,0.0020443148,0.00023265326,0.0007816719,0.0077891448,0.009913602,0.0032424962],"genre_scores_gemma":[0.91617376,0.00026915438,0.07186443,0.00030458494,0.00009898775,0.00041097013,0.009214507,0.00009387654,0.0015697143],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9985393,0.0008601664,0.00010081017,0.0002877681,0.00012680171,0.000085187414],"domain_scores_gemma":[0.9933878,0.0053554093,0.00036731077,0.0003333425,0.0004618545,0.00009422954],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034923856,0.0011324175,0.0005665144,0.0011387288,0.0002943773,0.0007483039,0.0010211597,0.00092947483,0.0016011794],"category_scores_gemma":[0.008533302,0.00032188243,0.0011079314,0.00062658073,0.00033215052,0.0009496951,0.0007303481,0.0014587903,0.00094662193],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003815901,0.0014732487,0.02978607,0.0009619527,0.000692526,0.0010166846,0.001033989,0.2745493,0.018808471,0.0022486288,0.013376564,0.6522367],"study_design_scores_gemma":[0.000050300772,0.00026473007,0.0037459978,0.000031655123,0.000083612766,0.00015063422,0.00008909509,0.9890676,0.0034948972,0.0017941949,0.0012027822,0.000024486553],"about_ca_topic_score_codex":0.0055136597,"about_ca_topic_score_gemma":0.007492333,"teacher_disagreement_score":0.0055136597,"about_ca_system_score_codex":0.0009923044,"about_ca_system_score_gemma":0.0010308356,"threshold_uncertainty_score":0.018469691},"labels":[],"label_agreement":null},{"id":"W4408135179","doi":"10.3389/fdgth.2025.1495040","title":"A simplified retriever to improve accuracy of phenotype normalizations by large language models","year":2025,"lang":"en","type":"article","venue":"Frontiers in Digital Health","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Normalization (sociology); Computer science; Natural language processing; Artificial intelligence; Labrador Retriever; Term (time); Language model; Machine learning; Data mining","score_opus":0.007173059222425137,"score_gpt":0.2910441178876113,"score_spread":0.2838710586651862,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408135179","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04960695,0.0011160508,0.86746675,0.00089911657,0.00024371964,0.00044687945,0.0062543876,0.06991034,0.004055845],"genre_scores_gemma":[0.19168189,0.0007272126,0.7774184,0.0007228137,0.00023172023,0.00049184443,0.018454818,0.003926233,0.0063450364],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972458,0.0007555132,0.00035677623,0.0006733045,0.0008631602,0.00010548151],"domain_scores_gemma":[0.9933329,0.00340469,0.0004094535,0.0017195026,0.0010117096,0.000121710895],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034189865,0.0019565187,0.0014314095,0.0030991824,0.0005766321,0.0020114665,0.00198965,0.0010743641,0.009882067],"category_scores_gemma":[0.017762102,0.00059670885,0.0019744243,0.0021299862,0.0005731731,0.0045276987,0.002996106,0.0016013432,0.008724117],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00068167306,0.0003437647,0.0051366836,0.0011663402,0.00041418054,0.00063493784,0.0006533476,0.017166402,0.07908715,0.007319309,0.044143736,0.8432525],"study_design_scores_gemma":[0.0003656639,0.0005524046,0.0057486016,0.0001661916,0.0005937329,0.0026033807,0.00072529796,0.76920986,0.11494615,0.029869495,0.07493828,0.00028090033],"about_ca_topic_score_codex":0.004776557,"about_ca_topic_score_gemma":0.0067554982,"teacher_disagreement_score":0.009882067,"about_ca_system_score_codex":0.00073332095,"about_ca_system_score_gemma":0.0017004408,"threshold_uncertainty_score":0.033058822},"labels":[],"label_agreement":null},{"id":"W4408177747","doi":"10.1093/jamiaopen/ooaf010","title":"Semantic enrichment of Pomeranian health study data using LOINC and WHO-FIC terminology mapping principles","year":2025,"lang":"en","type":"article","venue":"JAMIA Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Terminology; Computer science; Information retrieval; Natural language processing; Linguistics; Philosophy","score_opus":0.13018531241328457,"score_gpt":0.3993251602189576,"score_spread":0.269139847805673,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408177747","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06338824,0.0019199455,0.81355226,0.008049939,0.00091076456,0.003803541,0.06586702,0.007289863,0.03521843],"genre_scores_gemma":[0.09610923,0.00096015225,0.84591657,0.001279769,0.00017215138,0.0026034005,0.049453497,0.0011991046,0.0023061517],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.98304784,0.008072326,0.0041104914,0.0021945806,0.002203628,0.00037116287],"domain_scores_gemma":[0.9433311,0.031603154,0.004641106,0.012371852,0.007059131,0.0009935488],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.027856158,0.0010028438,0.0008983005,0.0166048,0.0018951765,0.005005967,0.0019880447,0.0010389928,0.00505162],"category_scores_gemma":[0.07494619,0.0007153258,0.0017521197,0.011481378,0.0025634605,0.006356058,0.008512426,0.0019600014,0.0018163658],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012412007,0.00032197844,0.05294768,0.011853773,0.00046892551,0.0024312902,0.04696563,0.007100915,0.01778424,0.24661577,0.09532842,0.51694024],"study_design_scores_gemma":[0.0001383111,0.00014011265,0.024649907,0.004941603,0.00035682213,0.0012813052,0.010525711,0.012347595,0.016806843,0.10665631,0.8219325,0.00022293972],"about_ca_topic_score_codex":0.007844621,"about_ca_topic_score_gemma":0.007461431,"teacher_disagreement_score":0.027856158,"about_ca_system_score_codex":0.0035732447,"about_ca_system_score_gemma":0.010514823,"threshold_uncertainty_score":0.1473192},"labels":[],"label_agreement":null},{"id":"W4408187354","doi":"10.1093/genetics/iyaf027","title":"The Unified Phenotype Ontology : a framework for cross-species integrative phenomics","year":2025,"lang":"en","type":"article","venue":"Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":10,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Basic Energy Sciences; Takeda Canada; Biotechnology and Biological Sciences Research Council; Office of Science; National Institutes of Health; Norges Idrettshøgskole; GlaxoSmithKline foundation; National Human Genome Research Institute; Wellcome Trust; Eunice Kennedy Shriver National Institute of Child Health and Human Development; U.S. Department of Energy; European Bioinformatics Institute; Centers for Disease Control and Prevention; Sanofi Australia; Biogen; Celgene; Office of the Director","keywords":"Phenomics; Ontology; Phenome; Biology; Data integration; Phenotypic trait; Controlled vocabulary; Computational biology; Open Biomedical Ontologies; Representation (politics); Biological data; Data science; Organism; Computer science; Vocabulary; Phenotype; Ontology-based data integration; Bioinformatics; Genomics; Data mining; Information retrieval; Genetics; Genome; Ontology alignment; Semantic Web; Gene","score_opus":0.02122785670023918,"score_gpt":0.3328977548537405,"score_spread":0.3116698981535013,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408187354","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00075084495,0.00043111085,0.9891722,0.0011824612,0.00011260558,0.00034068452,0.0016328088,0.0030466767,0.003330501],"genre_scores_gemma":[0.012034113,0.0012771827,0.97602975,0.0006721098,0.000114077855,0.000909146,0.006605529,0.00078452553,0.0015734812],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99268687,0.0024149264,0.0013705066,0.0011118794,0.0020343068,0.00038158539],"domain_scores_gemma":[0.99250835,0.0027413426,0.00080261746,0.0019568927,0.0013109101,0.000679883],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.015205983,0.0017729915,0.0019123172,0.009435434,0.0033665935,0.009183098,0.0075395126,0.0030630424,0.0035012385],"category_scores_gemma":[0.014170586,0.0016287899,0.0054323203,0.009330367,0.0045518335,0.014795357,0.008928536,0.006689853,0.0021534427],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000063539475,0.00012579894,0.0012667374,0.0007749704,0.00018515403,0.00081384956,0.0022794488,0.0135513665,0.0025062582,0.8624665,0.01903761,0.0969287],"study_design_scores_gemma":[0.00003529661,0.000046541874,0.00083261676,0.000839999,0.00013211492,0.00071345834,0.00090833777,0.057897814,0.0017508827,0.5846664,0.35204715,0.00012947609],"about_ca_topic_score_codex":0.024485707,"about_ca_topic_score_gemma":0.02344528,"teacher_disagreement_score":0.024485707,"about_ca_system_score_codex":0.0048122755,"about_ca_system_score_gemma":0.011927665,"threshold_uncertainty_score":0.08041787},"labels":[],"label_agreement":null},{"id":"W4408236585","doi":"10.1017/rsm.2024.18","title":"Translating systematic searches in the APA PsycInfo database from Ovid to EBSCOhost: A tutorial based on a filter translation","year":2025,"lang":"en","type":"article","venue":"Research Synthesis Methods","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Victoria","funders":"","keywords":"PsycINFO; Computer science; MEDLINE; Translation (biology); Psychology","score_opus":0.2116884651036383,"score_gpt":0.5070427482358725,"score_spread":0.29535428313223416,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408236585","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015823394,0.090835266,0.7445378,0.05977615,0.009671093,0.02124685,0.010007844,0.010830852,0.051511876],"genre_scores_gemma":[0.0039556525,0.08001396,0.8484206,0.017025232,0.001929849,0.01991752,0.003419094,0.0030019726,0.022316108],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.96524084,0.022758136,0.0073513826,0.0010895752,0.003099607,0.00046034434],"domain_scores_gemma":[0.88435477,0.10004768,0.0038660616,0.0026523883,0.007937883,0.0011412285],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04335422,0.0027259625,0.0027651477,0.013545841,0.0017235309,0.0072578453,0.002817922,0.004792155,0.057968687],"category_scores_gemma":[0.1322678,0.0029495235,0.0034996385,0.016548442,0.003103908,0.013893544,0.0052244063,0.006288913,0.03506313],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020138199,0.000121243924,0.00023387467,0.028341917,0.00018458176,0.0005357895,0.00615213,0.001299113,0.0018691168,0.063264385,0.3720323,0.52576417],"study_design_scores_gemma":[0.00013163888,0.00018420084,0.0004767688,0.015081999,0.00005889917,0.00071645225,0.00086125836,0.0013986739,0.0006760127,0.03296282,0.94733703,0.00011424568],"about_ca_topic_score_codex":0.005264232,"about_ca_topic_score_gemma":0.008653405,"teacher_disagreement_score":0.9566458,"about_ca_system_score_codex":0.004905026,"about_ca_system_score_gemma":0.011758406,"threshold_uncertainty_score":0.22928184},"labels":[],"label_agreement":null},{"id":"W4408531621","doi":"10.1016/j.gimo.2025.102864","title":"P020: Introducing an interactive, searchable database of LC-FAOD gene variants, genotypes and phenotypes","year":2025,"lang":"en","type":"article","venue":"Genetics in Medicine Open","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Phenotype; Genotype; Gene; Genetics; Genotype-phenotype distinction; Biology; Database; Computational biology; Bioinformatics; Computer science","score_opus":0.027282918863029124,"score_gpt":0.36060083665687964,"score_spread":0.3333179177938505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408531621","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020658098,0.0061170845,0.39325422,0.009802421,0.001259785,0.0020341366,0.37832004,0.12419178,0.06436248],"genre_scores_gemma":[0.1466733,0.0073654465,0.44170627,0.00582111,0.00065892056,0.0037538777,0.3607884,0.011195835,0.022036826],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988133,0.00037421266,0.00021566649,0.00021313396,0.00032322464,0.00006044615],"domain_scores_gemma":[0.9942095,0.0038834417,0.00040617643,0.0006092937,0.0004957147,0.00039587007],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0038027598,0.001011248,0.00085039064,0.0030621674,0.00040152867,0.0028234143,0.0015684398,0.0015993675,0.04835215],"category_scores_gemma":[0.014751661,0.00048978545,0.001052514,0.0018824914,0.00033992474,0.0020055824,0.00339081,0.0011817646,0.015154495],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0029706445,0.00047304656,0.009442418,0.0041581676,0.0006177323,0.0016780556,0.00036647648,0.0057383343,0.010991721,0.014915414,0.5474973,0.40115076],"study_design_scores_gemma":[0.0010645094,0.00035802744,0.006319009,0.00091409986,0.0004633093,0.0021236169,0.00022816502,0.0251473,0.016118553,0.044413164,0.90269905,0.00015130154],"about_ca_topic_score_codex":0.0016930084,"about_ca_topic_score_gemma":0.0037318678,"teacher_disagreement_score":0.04835215,"about_ca_system_score_codex":0.00067670457,"about_ca_system_score_gemma":0.001856118,"threshold_uncertainty_score":0.16175407},"labels":[],"label_agreement":null},{"id":"W4408661451","doi":"10.63485/dwpwg-5wm24","title":"Notes on OA to PSI in Canada","year":2009,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.019122195524997814,"score_gpt":0.27329315270393034,"score_spread":0.2541709571789325,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408661451","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001584353,0.01099305,0.001410681,0.15863945,0.027675048,0.00024812104,0.013574048,0.0014839995,0.78439116],"genre_scores_gemma":[0.0070327385,0.004946635,0.0007597513,0.013040201,0.0017424315,0.000062590865,0.0036186196,0.0005660036,0.96823096],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99491054,0.00015965263,0.00013992455,0.00040705633,0.0031701718,0.0012127345],"domain_scores_gemma":[0.9916747,0.0006954446,0.00009468862,0.00047520525,0.0052761654,0.0017838017],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.0029429176,0.00084307167,0.0009640198,0.0028501237,0.010068848,0.008513577,0.001794172,0.0034637214,0.2767208],"category_scores_gemma":[0.008496845,0.0005024169,0.0008859074,0.0073764636,0.001756345,0.0041881544,0.003981761,0.0041693416,0.069181986],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011229285,0.000006451437,0.00007750556,0.000015651412,0.0000013127617,0.000040251136,0.00007742326,0.0000250173,0.000042347485,0.0037515503,0.989047,0.006904179],"study_design_scores_gemma":[0.0000038393187,0.0000024136698,0.00044663306,0.000025442296,0.0000012764374,0.00000976367,0.00015467564,0.000019906758,0.000036058893,0.00034654426,0.998946,0.000007319321],"about_ca_topic_score_codex":0.9484039,"about_ca_topic_score_gemma":0.9665391,"teacher_disagreement_score":0.99820584,"about_ca_system_score_codex":0.036090497,"about_ca_system_score_gemma":0.073096864,"threshold_uncertainty_score":0.92572325},"labels":[],"label_agreement":null},{"id":"W4408669907","doi":"10.63485/367xn-hhq70","title":"More on OA at U. Calgary","year":2009,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Medicine","score_opus":0.01968139266621793,"score_gpt":0.29384188345366696,"score_spread":0.27416049078744903,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408669907","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00027442185,0.0065412587,0.00086656277,0.017772876,0.016721684,0.00016587242,0.016602239,0.003434948,0.93762016],"genre_scores_gemma":[0.0011165048,0.0018870914,0.00044298085,0.005492928,0.0017667945,0.00007385214,0.0027192354,0.001365077,0.9851355],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99898666,0.000090440975,0.000062931125,0.00019526198,0.00051750656,0.00014723736],"domain_scores_gemma":[0.99590105,0.00046825557,0.0001723127,0.0006376133,0.0017886333,0.0010321935],"candidate_categories":["open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011100692,0.0010197185,0.0015906631,0.0064010695,0.00419869,0.008824479,0.0014882755,0.0029761798,0.84797543],"category_scores_gemma":[0.0057111094,0.0006119889,0.00067902124,0.009792223,0.00096712745,0.0051036426,0.004030462,0.0032320528,0.68582916],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000009401816,0.000008775199,0.000049645238,0.000032749125,0.0000014272616,0.000025203504,0.000032889813,0.000010270573,0.000043943073,0.00083642465,0.9818226,0.017126618],"study_design_scores_gemma":[0.000003728894,0.000001965034,0.00030911787,0.00005598958,8.435353e-7,0.0000077452605,0.000035057572,0.000006695042,0.000018724015,0.00029198395,0.9992644,0.0000037348693],"about_ca_topic_score_codex":0.08608562,"about_ca_topic_score_gemma":0.18109551,"teacher_disagreement_score":0.99851173,"about_ca_system_score_codex":0.0047461186,"about_ca_system_score_gemma":0.003696042,"threshold_uncertainty_score":0.21684462},"labels":[],"label_agreement":null},{"id":"W4408712380","doi":"10.63485/ykt8p-qvg55","title":"Call for OA to publicly-funded research in Canada","year":2009,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Business","score_opus":0.11372730319028909,"score_gpt":0.3927896886517479,"score_spread":0.2790623854614588,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408712380","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002002285,0.008514977,0.0018704707,0.22296,0.04267932,0.0013200704,0.047141157,0.007297825,0.66621387],"genre_scores_gemma":[0.0060355677,0.0029619557,0.0021688575,0.030091895,0.0021497773,0.00044545627,0.011153243,0.0013369928,0.9436564],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99209994,0.0003528463,0.00027089939,0.0006766481,0.004360242,0.0022393344],"domain_scores_gemma":[0.9542878,0.0017047323,0.00076065643,0.002513066,0.02068417,0.020049514],"candidate_categories":["scholarly_communication","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006509602,0.0012638458,0.0017112137,0.0061612027,0.009150221,0.012385032,0.0036005036,0.008216234,0.49883392],"category_scores_gemma":[0.01833794,0.0008705955,0.0019838845,0.0076218015,0.0027765264,0.004633312,0.009569269,0.0034728732,0.24661061],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000033364868,0.000014506927,0.00022361535,0.00006418459,0.0000053737413,0.00007489004,0.000025419207,0.000015424868,0.00008590994,0.0018044835,0.9892146,0.008438234],"study_design_scores_gemma":[0.000027371123,0.000007976444,0.0014576634,0.000088088615,0.0000049603486,0.000024700434,0.00016955932,0.000030513502,0.00004880366,0.00049835426,0.99762493,0.00001708896],"about_ca_topic_score_codex":0.8327561,"about_ca_topic_score_gemma":0.90271354,"teacher_disagreement_score":0.9963995,"about_ca_system_score_codex":0.04216215,"about_ca_system_score_gemma":0.1865094,"threshold_uncertainty_score":0.7148526},"labels":[],"label_agreement":null},{"id":"W4408766524","doi":"10.63485/h6xmx-4qf96","title":"More on OA to publicly-funded research in Canada","year":2008,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Regional science; Library science; Computer science; Geography","score_opus":0.1119213212652137,"score_gpt":0.37688502514139616,"score_spread":0.26496370387618245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408766524","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0062366324,0.03217689,0.0029836264,0.30369663,0.01479398,0.00022053784,0.020874068,0.0015279467,0.61748976],"genre_scores_gemma":[0.08278962,0.030007955,0.0064006513,0.13312715,0.005288437,0.00030649258,0.016599152,0.0021850741,0.72329545],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.97466284,0.0018167268,0.0008631068,0.0017412173,0.014063688,0.0068525122],"domain_scores_gemma":[0.9373668,0.009813476,0.0014594415,0.004647529,0.032669153,0.01404356],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.012558069,0.0011164561,0.0016092835,0.01668098,0.015561205,0.024727317,0.0029561582,0.0057333256,0.14395648],"category_scores_gemma":[0.03576653,0.0007672548,0.0018447058,0.040811807,0.0062679835,0.011422651,0.00961593,0.0052756076,0.016723974],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048076057,0.000037941052,0.001410447,0.00023539193,0.000028548679,0.0002645598,0.0015576783,0.00014641904,0.00018137884,0.054565288,0.9084252,0.033099167],"study_design_scores_gemma":[0.000019501409,0.0000066998896,0.0042463215,0.0003008184,0.0000131596225,0.00003570837,0.0014898899,0.000072964096,0.00006063219,0.0034984364,0.9902104,0.00004544486],"about_ca_topic_score_codex":0.97811913,"about_ca_topic_score_gemma":0.9874467,"teacher_disagreement_score":0.99704385,"about_ca_system_score_codex":0.13393782,"about_ca_system_score_gemma":0.18714522,"threshold_uncertainty_score":0.97179145},"labels":[],"label_agreement":null},{"id":"W4408784882","doi":"10.63485/gjjkv-98r98","title":"Author attitudes toward SPARC Partner journals","year":2008,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.11757209753796448,"score_gpt":0.3813945265472326,"score_spread":0.2638224290092681,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408784882","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9412086,0.0013517347,0.00026149765,0.01131689,0.0003320259,0.000052052972,0.00091307086,0.000038794944,0.044525478],"genre_scores_gemma":[0.9805058,0.0011999843,0.00020818842,0.0010837927,0.0001481351,0.000014429338,0.00031182935,0.00004540499,0.016482322],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9883536,0.0020508426,0.00061899,0.000759826,0.006924945,0.0012917245],"domain_scores_gemma":[0.8529391,0.041351017,0.024627483,0.0032653722,0.0489602,0.028856687],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.007814339,0.00020328621,0.00030292518,0.0070089204,0.005230556,0.0072085112,0.00063060364,0.0009380369,0.014926627],"category_scores_gemma":[0.07640805,0.00023841915,0.000198947,0.009931082,0.0019950487,0.0018291022,0.0019288952,0.001011643,0.0017016813],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00084814086,0.0002482153,0.7646719,0.00024309795,0.00009650441,0.0011224751,0.104841016,0.00025781745,0.0020997424,0.008969317,0.03968401,0.076917626],"study_design_scores_gemma":[0.00009212116,0.00023581888,0.73650426,0.00035995626,0.0000746857,0.0010777925,0.11835054,0.00035929156,0.0014173031,0.001853724,0.13948643,0.00018803956],"about_ca_topic_score_codex":0.34242347,"about_ca_topic_score_gemma":0.43869013,"teacher_disagreement_score":0.9927915,"about_ca_system_score_codex":0.009664743,"about_ca_system_score_gemma":0.01006164,"threshold_uncertainty_score":0.6808607},"labels":[],"label_agreement":null},{"id":"W4408877701","doi":"10.63485/e4rqn-7h881","title":"APLA converts its journal to OA","year":2007,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science","score_opus":0.031995989513365374,"score_gpt":0.325892348883806,"score_spread":0.2938963593704407,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408877701","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020291246,0.0019991698,0.0061196676,0.068135016,0.2292017,0.0004568501,0.010293363,0.012968081,0.66879696],"genre_scores_gemma":[0.0068000513,0.0014085291,0.00252456,0.015169681,0.029967738,0.00021282773,0.007316488,0.0050919103,0.9315081],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9920367,0.00063746155,0.0005231852,0.00067554467,0.005284699,0.0008423879],"domain_scores_gemma":[0.97059894,0.004579222,0.00092961214,0.0040196776,0.01470039,0.0051721893],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.0057053016,0.0008223836,0.00089653814,0.0060050385,0.0038580967,0.02055875,0.0016931703,0.0026509736,0.1958841],"category_scores_gemma":[0.038591653,0.0006119792,0.00079856056,0.004520977,0.002608094,0.0054884,0.0048273685,0.006055662,0.17695932],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015298856,0.000011573662,0.000049716007,0.000023187833,0.0000017568995,0.000042253905,0.000049389455,0.000010272039,0.0001113992,0.0035608208,0.9861925,0.00993172],"study_design_scores_gemma":[0.0000039998818,0.0000035544351,0.00009544752,0.000013184482,7.8179306e-7,0.000012740944,0.000040315976,0.00001582236,0.0000656805,0.00050290476,0.9992424,0.000003240986],"about_ca_topic_score_codex":0.02792066,"about_ca_topic_score_gemma":0.032027487,"teacher_disagreement_score":0.9983068,"about_ca_system_score_codex":0.006351615,"about_ca_system_score_gemma":0.009489742,"threshold_uncertainty_score":0.6552976},"labels":[],"label_agreement":null},{"id":"W4408877965","doi":"10.63485/9d10j-nbx03","title":"Canada is missing an opportunity for OA to publicly-funded research (and why)","year":2007,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science; Library science; Computer science","score_opus":0.22205321279140666,"score_gpt":0.4231289723823648,"score_spread":0.20107575959095814,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408877965","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011221687,0.0036997928,0.00044388874,0.92457175,0.011343439,0.000041643034,0.00094512664,0.00029340596,0.057538766],"genre_scores_gemma":[0.025860894,0.0053373706,0.0019716863,0.6827858,0.0035412896,0.00008375157,0.00084055995,0.00047985232,0.2790989],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9842517,0.0009249529,0.00049105193,0.0011881089,0.009519976,0.0036242655],"domain_scores_gemma":[0.961552,0.0073388442,0.0010322633,0.002220225,0.01777709,0.010079539],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.0110771265,0.00059270655,0.0011394754,0.0018065433,0.017393194,0.020932883,0.0022693346,0.019008929,0.045637775],"category_scores_gemma":[0.031462613,0.00074660406,0.0011333245,0.004327397,0.009988138,0.008142334,0.0036104165,0.012845632,0.015052798],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018652481,0.000009203782,0.00033017097,0.000042364376,0.000008540183,0.00011425234,0.00031555953,0.000024726292,0.00013157802,0.022102538,0.9703799,0.0065225363],"study_design_scores_gemma":[0.000017453562,0.0000053183617,0.0010829994,0.00007510947,0.00001100248,0.000032114167,0.00078394584,0.00003667399,0.00013292787,0.0030367686,0.99474686,0.000038753264],"about_ca_topic_score_codex":0.9231813,"about_ca_topic_score_gemma":0.9602362,"teacher_disagreement_score":0.9977307,"about_ca_system_score_codex":0.043104246,"about_ca_system_score_gemma":0.17288546,"threshold_uncertainty_score":0.31274468},"labels":[],"label_agreement":null},{"id":"W4408911393","doi":"10.1007/s00371-025-03857-1","title":"A visual-language foundation model for disease diagnosis and doctor–patient co-decision","year":2025,"lang":"en","type":"article","venue":"The Visual Computer","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Foundation (evidence); Disease; Computer science; Medicine; Artificial intelligence; Pathology","score_opus":0.017516799662783227,"score_gpt":0.36238848785635563,"score_spread":0.3448716881935724,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4408911393","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0046326225,0.00014692925,0.9858807,0.00094353437,0.000057594592,0.00016764762,0.0020040527,0.0032006763,0.0029662484],"genre_scores_gemma":[0.22751217,0.0003503418,0.76104504,0.000551928,0.00006485396,0.00041628967,0.004563499,0.00037433387,0.0051216492],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986338,0.00040251997,0.00017116318,0.00028053054,0.0003897765,0.00012216439],"domain_scores_gemma":[0.99648833,0.0016759711,0.00022942133,0.0005658274,0.0008244047,0.0002161325],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0025674358,0.00055526453,0.00055211096,0.0018286,0.0006193025,0.0027922143,0.0018536187,0.0012053599,0.006443632],"category_scores_gemma":[0.008572832,0.00045207058,0.0017318849,0.001414121,0.0007694858,0.0038004355,0.001842135,0.0012570681,0.0018329342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082129706,0.0004898491,0.0051775817,0.000635779,0.00026742124,0.00078478584,0.00094774575,0.15438886,0.0061079124,0.46986854,0.03620291,0.3243074],"study_design_scores_gemma":[0.00006213798,0.00005293545,0.00046866262,0.00011457045,0.00009153293,0.0001940974,0.00011233185,0.80063444,0.0038037274,0.1676468,0.026784068,0.00003468567],"about_ca_topic_score_codex":0.02088348,"about_ca_topic_score_gemma":0.02681786,"teacher_disagreement_score":0.02088348,"about_ca_system_score_codex":0.0017820148,"about_ca_system_score_gemma":0.0031877824,"threshold_uncertainty_score":0.041523814},"labels":[],"label_agreement":null},{"id":"W4409080179","doi":"10.59350/x63b5-jwm28","title":"Montreal BioJava Bootcamp Announced","year":2003,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science","score_opus":0.01773147944329089,"score_gpt":0.2728545852523451,"score_spread":0.25512310580905423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409080179","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036802893,0.009119755,0.0059619183,0.05673072,0.024681062,0.00048921414,0.021508547,0.0064460137,0.87138253],"genre_scores_gemma":[0.004419589,0.0011650297,0.0019153828,0.0016786918,0.0006670307,0.000069953654,0.004895497,0.00069315446,0.9844956],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99823403,0.00011003572,0.00002129286,0.00032772182,0.0009276476,0.0003793259],"domain_scores_gemma":[0.9961422,0.00017228689,0.00006347365,0.00032097643,0.0018496108,0.00145143],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0023972222,0.0012933458,0.00082023407,0.0016917225,0.006286312,0.0067901476,0.0016373192,0.0023411813,0.4520676],"category_scores_gemma":[0.0033554477,0.00072993967,0.0008095186,0.0018627021,0.0009988877,0.0021242443,0.0026697323,0.0033474981,0.17376928],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040015162,0.0000133489175,0.000108740955,0.0000127284575,0.0000018562614,0.000067486486,0.000026690668,0.000030974443,0.000359851,0.004002925,0.9851841,0.010151279],"study_design_scores_gemma":[0.00000646618,0.000005561648,0.00025028302,0.000007843518,8.491422e-7,0.000017737748,0.000023796041,0.000044548022,0.00013918635,0.00032564145,0.9991736,0.000004457256],"about_ca_topic_score_codex":0.5074538,"about_ca_topic_score_gemma":0.7233091,"teacher_disagreement_score":0.4925462,"about_ca_system_score_codex":0.011720129,"about_ca_system_score_gemma":0.016964281,"threshold_uncertainty_score":0.99089384},"labels":[],"label_agreement":null},{"id":"W4409227183","doi":"10.1212/wnl.0000000000210492","title":"Application of a Clinical Algorithm on Real-World Electronic Medical Record Data to Assist With Earlier Detection of Amyotrophic Lateral Sclerosis (P8-2.010)","year":2025,"lang":"en","type":"article","venue":"Neurology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Montreal Neurological Institute and Hospital; Trillium Health Centre","funders":"","keywords":"Amyotrophic lateral sclerosis; Medical record; Computer science; Medicine; Artificial intelligence; Algorithm; Pathology; Disease; Internal medicine","score_opus":0.03542111636281635,"score_gpt":0.3387313170713792,"score_spread":0.30331020070856285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409227183","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.49070445,0.0021993301,0.44209826,0.012092296,0.00073997333,0.0033149153,0.023074172,0.00988045,0.015896123],"genre_scores_gemma":[0.5197209,0.0005180761,0.4638366,0.00094150443,0.00009482984,0.00054179755,0.01259331,0.00009755561,0.0016554532],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99876934,0.00034322625,0.00022899454,0.0003611265,0.0002508751,0.000046483856],"domain_scores_gemma":[0.99528813,0.0024324912,0.00060429965,0.00025525846,0.0012544626,0.00016537817],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024823332,0.00037309589,0.00037117096,0.0023640685,0.0004511874,0.0015150699,0.00048004734,0.0007548857,0.0014334325],"category_scores_gemma":[0.011073351,0.00013994411,0.00048093632,0.0015238894,0.00018050292,0.00078024593,0.0006374081,0.00049800344,0.00066145975],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001203011,0.001723703,0.20036218,0.00082236214,0.00076489587,0.0018173406,0.0005581489,0.030184664,0.01328305,0.0054328544,0.04111961,0.7027282],"study_design_scores_gemma":[0.000778827,0.00085145456,0.11720389,0.00073786004,0.0009854215,0.0039800066,0.00088606984,0.77216953,0.025550324,0.016322816,0.060413297,0.000120557495],"about_ca_topic_score_codex":0.00489731,"about_ca_topic_score_gemma":0.008559305,"teacher_disagreement_score":0.00489731,"about_ca_system_score_codex":0.00087494444,"about_ca_system_score_gemma":0.0026380233,"threshold_uncertainty_score":0.013128042},"labels":[],"label_agreement":null},{"id":"W4409457086","doi":"10.3390/app15084353","title":"Developing Contextual Ontology for Chronic Diseases: AI-Enhanced Extension and Prediction in an Asthma Case Study","year":2025,"lang":"en","type":"article","venue":"Applied Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université du Québec à Rimouski; Université du Québec à Chicoutimi","funders":"","keywords":"Asthma; Extension (predicate logic); Ontology; Computer science; Medicine; Data science; Immunology; Epistemology; Philosophy; Programming language","score_opus":0.028292125074967932,"score_gpt":0.34792732788468567,"score_spread":0.31963520280971774,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409457086","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.47304294,0.0012794611,0.49222952,0.0064095072,0.00016662519,0.0012866178,0.007493183,0.0011391722,0.01695302],"genre_scores_gemma":[0.6119474,0.0007523514,0.38108477,0.00036026214,0.000035990553,0.00038721913,0.0041314894,0.00006982288,0.0012307271],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99824905,0.0007487473,0.0002167187,0.0002487395,0.0004027757,0.00013392766],"domain_scores_gemma":[0.9962775,0.0025873452,0.00019484578,0.00034363713,0.00047317124,0.0001234766],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0029551506,0.00040355232,0.000348643,0.001667949,0.0007881784,0.0014715666,0.000793444,0.0010675086,0.0014583943],"category_scores_gemma":[0.00847639,0.00018001822,0.0010076646,0.0016820248,0.0007088779,0.0014975473,0.0013927129,0.00083897775,0.00015149023],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008776834,0.0015237009,0.17143996,0.0016690765,0.00040621657,0.021352287,0.008215894,0.19622077,0.017655943,0.112496324,0.017845822,0.4502963],"study_design_scores_gemma":[0.0001844839,0.00025292934,0.033235103,0.0006120934,0.00038227736,0.005406429,0.005163497,0.78842664,0.013831387,0.0546465,0.09772335,0.00013533],"about_ca_topic_score_codex":0.026990354,"about_ca_topic_score_gemma":0.03503497,"teacher_disagreement_score":0.026990354,"about_ca_system_score_codex":0.0015163018,"about_ca_system_score_gemma":0.0027947368,"threshold_uncertainty_score":0.053666532},"labels":[],"label_agreement":null},{"id":"W4409487172","doi":"10.63485/a33qx-2rz35","title":"Another call for OA to Canadian research","year":2007,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Political science","score_opus":0.13045832088929363,"score_gpt":0.42728787428900805,"score_spread":0.2968295533997144,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409487172","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00041873375,0.0035968486,0.0029695532,0.85061496,0.04067006,0.00016534855,0.001328027,0.0009601604,0.09927631],"genre_scores_gemma":[0.011860203,0.0052091638,0.01809084,0.5283291,0.009254657,0.00032209902,0.0018284299,0.001996489,0.42310908],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.94744706,0.0071239313,0.0024004423,0.0057715047,0.026530242,0.010726887],"domain_scores_gemma":[0.7603965,0.018997911,0.003968566,0.024514407,0.12336402,0.06875862],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.052779935,0.0014647914,0.0028488624,0.008150372,0.016515376,0.03176517,0.006900008,0.026324024,0.18054432],"category_scores_gemma":[0.122697726,0.0014568134,0.0035768885,0.008533838,0.0146044595,0.025781434,0.020645844,0.03028069,0.0631378],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007351961,0.000033193424,0.0003406129,0.00013721197,0.000025225798,0.00012539714,0.0002719646,0.000045263976,0.00024587396,0.07501064,0.9103344,0.013356657],"study_design_scores_gemma":[0.000036900616,0.000009001361,0.00038806803,0.00015071158,0.000011371941,0.000046811147,0.00046411657,0.00004680141,0.00006688755,0.008703029,0.9900456,0.000030706557],"about_ca_topic_score_codex":0.5959008,"about_ca_topic_score_gemma":0.68102366,"teacher_disagreement_score":0.9931,"about_ca_system_score_codex":0.044414137,"about_ca_system_score_gemma":0.20579556,"threshold_uncertainty_score":0.8129581},"labels":[],"label_agreement":null},{"id":"W4409534569","doi":"10.1016/j.parkreldis.2025.107395","title":"Demonstration of a Prototype Clinical Decision Support System for Diagnosing Parkinson’s Disease","year":2025,"lang":"en","type":"article","venue":"Parkinsonism & Related Disorders","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre for Movement Disorders","funders":"Case Western Reserve University; U.S. Department of Defense","keywords":"Parkinson's disease; Disease; Decision support system; Medicine; Clinical decision support system; Physical medicine and rehabilitation; Computer science; Neuroscience; Psychology; Artificial intelligence; Pathology","score_opus":0.01229901224252512,"score_gpt":0.3143200285122731,"score_spread":0.302021016269748,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409534569","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.22622843,0.0016781136,0.5997476,0.00841671,0.001248539,0.002387243,0.012343617,0.12732667,0.020623107],"genre_scores_gemma":[0.55092204,0.000528032,0.42714146,0.0016757699,0.00015906586,0.0005974042,0.006726604,0.0010063449,0.01124331],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995863,0.000089956695,0.00006271808,0.000108550616,0.00012270309,0.000029760033],"domain_scores_gemma":[0.99792874,0.0013034068,0.00006767181,0.00016620115,0.00034487323,0.00018905777],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010333931,0.00044187164,0.00046385167,0.0006288931,0.00040471158,0.0012790249,0.0009632279,0.0009154432,0.010415861],"category_scores_gemma":[0.0037441168,0.00027380727,0.000333072,0.00048149042,0.00019489136,0.0007703475,0.0006473221,0.0006051761,0.0030699521],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0036368144,0.0015145482,0.022880461,0.0014654775,0.00031103182,0.008669668,0.0014670052,0.012587675,0.09990847,0.0047748527,0.13189988,0.71088403],"study_design_scores_gemma":[0.0024740049,0.0016492514,0.031184698,0.0006156958,0.00047985613,0.011229859,0.0012376637,0.5591225,0.1595371,0.011497038,0.22070251,0.0002698331],"about_ca_topic_score_codex":0.0037672815,"about_ca_topic_score_gemma":0.0035880115,"teacher_disagreement_score":0.010415861,"about_ca_system_score_codex":0.00038119341,"about_ca_system_score_gemma":0.0009792855,"threshold_uncertainty_score":0.034844518},"labels":[],"label_agreement":null},{"id":"W4409544291","doi":"10.63485/43egp-qck56","title":"An OA repository for community health in British Columbia","year":2006,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Community health; Medicine; Nursing; Public health","score_opus":0.02087917231675837,"score_gpt":0.30499444953572924,"score_spread":0.2841152772189709,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409544291","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03870887,0.017126773,0.023438988,0.02340389,0.0015329852,0.0013166124,0.46020517,0.02328302,0.41098368],"genre_scores_gemma":[0.19644025,0.018045273,0.08174662,0.00320407,0.00041805333,0.0008708257,0.33748266,0.006625528,0.35516667],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983943,0.00023962246,0.00016782273,0.0002002029,0.00075837574,0.00023960487],"domain_scores_gemma":[0.9846259,0.0018616195,0.00057730445,0.0020493444,0.007350163,0.0035357126],"candidate_categories":["scholarly_communication","open_science"],"consensus_categories":[],"category_scores_codex":[0.002647053,0.00046090852,0.0006974067,0.010800705,0.0048408844,0.00785706,0.001626192,0.0010308577,0.052502412],"category_scores_gemma":[0.012941763,0.00044707162,0.000323558,0.030707661,0.0011566228,0.002846114,0.0029982592,0.0011184814,0.012929835],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024386669,0.000069988884,0.0063784346,0.0007935406,0.00003085442,0.0005948817,0.002303797,0.0005327209,0.000956365,0.016203828,0.7402269,0.23166487],"study_design_scores_gemma":[0.000038695012,0.000012674454,0.0124293035,0.00049567607,0.000026247608,0.00012718157,0.0014950888,0.00094733114,0.00042403,0.0035563766,0.9803925,0.00005481323],"about_ca_topic_score_codex":0.8668924,"about_ca_topic_score_gemma":0.8948736,"teacher_disagreement_score":0.9983738,"about_ca_system_score_codex":0.016027072,"about_ca_system_score_gemma":0.06974168,"threshold_uncertainty_score":0.267783},"labels":[],"label_agreement":null},{"id":"W4409591038","doi":"10.63485/cx67j-45b60","title":"OA and peer review","year":2006,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Peer review; Computer science; Political science; Law","score_opus":0.02844522741997957,"score_gpt":0.3152159470496463,"score_spread":0.2867707196296667,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409591038","genre_codex":"editorial","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0024150857,0.01774829,0.03803674,0.15568966,0.4626547,0.009993974,0.012202211,0.012144694,0.28911468],"genre_scores_gemma":[0.018700289,0.01280013,0.02390124,0.014857527,0.10907452,0.0053754565,0.018008934,0.008036278,0.7892456],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.890018,0.020619897,0.015533864,0.0103491405,0.05589978,0.007579345],"domain_scores_gemma":[0.40967137,0.026836766,0.01623573,0.094244674,0.4148883,0.03812319],"candidate_categories":["metaresearch","scholarly_communication","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.064149685,0.0022915788,0.0059650536,0.015449316,0.0074387197,0.031787984,0.0061096274,0.008884962,0.36107826],"category_scores_gemma":[0.29881766,0.0018148014,0.0034481597,0.007883622,0.006684288,0.014546922,0.0130745135,0.009327066,0.42026895],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000096868214,0.000038866325,0.0002617423,0.0011380969,0.0000575659,0.0002478658,0.00027360316,0.000042556603,0.0007975286,0.0040975497,0.9191685,0.07377929],"study_design_scores_gemma":[0.000032863118,0.000020160753,0.00055181753,0.0007547309,0.000026161684,0.00025413412,0.00022956148,0.00011364985,0.00029286608,0.00395265,0.9937356,0.00003590103],"about_ca_topic_score_codex":0.0017924025,"about_ca_topic_score_gemma":0.002985844,"teacher_disagreement_score":0.99389035,"about_ca_system_score_codex":0.005479503,"about_ca_system_score_gemma":0.03705735,"threshold_uncertainty_score":0.9113443},"labels":[],"label_agreement":null},{"id":"W4409603033","doi":"10.1101/2025.04.18.25326076","title":"AutoReporter: Development of an artificial intelligence tool for automated assessment of research reporting guideline adherence","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; McMaster University; University of Toronto","funders":"","keywords":"Guideline; Computer science; Medicine","score_opus":0.19076275979764373,"score_gpt":0.5019322412940268,"score_spread":0.31116948149638307,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409603033","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015373903,0.0016863297,0.3460979,0.0033042561,0.0005848082,0.0036999122,0.07739601,0.5461719,0.0056849946],"genre_scores_gemma":[0.059835482,0.0007014542,0.83674055,0.0018692286,0.00018585526,0.0043973043,0.083554275,0.008602614,0.004113152],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9786339,0.009982763,0.0043261615,0.0033320985,0.003397488,0.00032757293],"domain_scores_gemma":[0.90715057,0.062086303,0.009514914,0.009534247,0.0104803685,0.0012336286],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.031708244,0.00247097,0.0016742678,0.0063323416,0.0007578802,0.0042128204,0.0031857511,0.0017105932,0.019398738],"category_scores_gemma":[0.12710445,0.0014328259,0.0022062138,0.0027844328,0.00052519346,0.003859831,0.0044942154,0.002334092,0.012178009],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001690296,0.00045187064,0.010874765,0.009045048,0.0007624477,0.0004336117,0.0017819072,0.0095381485,0.013027465,0.0076766927,0.41248417,0.53223354],"study_design_scores_gemma":[0.0020427112,0.00097457087,0.0149948215,0.0038830258,0.00089082174,0.0009344795,0.0012780863,0.3867401,0.071861394,0.037191644,0.47854736,0.0006608706],"about_ca_topic_score_codex":0.0041618785,"about_ca_topic_score_gemma":0.008682119,"teacher_disagreement_score":0.96829176,"about_ca_system_score_codex":0.0022370736,"about_ca_system_score_gemma":0.0071565486,"threshold_uncertainty_score":0.16769123},"labels":[],"label_agreement":null},{"id":"W4409616225","doi":"10.63485/w5anm-wn751","title":"OCA book-scanning at the U of Toronto","year":2005,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Art history; Art; Sociology; Media studies","score_opus":0.01648292617327455,"score_gpt":0.288065033315862,"score_spread":0.2715821071425874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409616225","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018543076,0.013495091,0.016471926,0.00575589,0.002173623,0.00025011963,0.28428453,0.05789111,0.6178233],"genre_scores_gemma":[0.0078088553,0.009261061,0.02491104,0.0010115968,0.00047282202,0.00022589194,0.13041127,0.017030824,0.8088666],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987821,0.00005973152,0.00007636543,0.00023386795,0.0007437966,0.000104171835],"domain_scores_gemma":[0.99620026,0.00055246125,0.00019594775,0.00085810066,0.0017134049,0.00047978025],"candidate_categories":["scholarly_communication","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012060548,0.0018446792,0.0014511822,0.006835942,0.0034818952,0.007986261,0.0014601975,0.0015105632,0.52800995],"category_scores_gemma":[0.0059197275,0.0013232462,0.00084582315,0.015914029,0.0011822325,0.0045250775,0.0025842069,0.0013116654,0.3657517],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015343621,0.0000032152932,0.00010974176,0.00013016642,0.0000040250925,0.000044477107,0.0001198068,0.00003332923,0.00027708468,0.0017694852,0.959595,0.03789836],"study_design_scores_gemma":[0.0000049470204,0.0000023328168,0.001109877,0.00007834717,0.000005171081,0.000057892514,0.00009247969,0.000080875434,0.00041627034,0.0008118543,0.99732697,0.000013070906],"about_ca_topic_score_codex":0.2982853,"about_ca_topic_score_gemma":0.5574158,"teacher_disagreement_score":0.9985398,"about_ca_system_score_codex":0.00795952,"about_ca_system_score_gemma":0.010279027,"threshold_uncertainty_score":0.6732365},"labels":[],"label_agreement":null},{"id":"W4409643066","doi":"10.63485/j0cvb-d9521","title":"OA to Canadian research data","year":2005,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Data science","score_opus":0.27381168888196233,"score_gpt":0.46375649800976393,"score_spread":0.1899448091278016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409643066","genre_codex":"dataset","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00029067448,0.00072792673,0.001092352,0.00175763,0.00033234427,0.00016340322,0.9293476,0.0016277573,0.06466035],"genre_scores_gemma":[0.004160042,0.0023092574,0.0067332904,0.0008723428,0.00011975323,0.00049252575,0.9433918,0.0009929326,0.040928118],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98738575,0.00090182293,0.001509663,0.0017074278,0.006736609,0.0017586338],"domain_scores_gemma":[0.9396574,0.004381494,0.0020207085,0.010584325,0.03935294,0.004003181],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.008040494,0.0018563521,0.0019699952,0.031765316,0.005237204,0.012995911,0.0036404221,0.0018324562,0.19049878],"category_scores_gemma":[0.055767547,0.0012675312,0.0018415608,0.06282398,0.0012434302,0.0033343295,0.006126653,0.0027382986,0.11085968],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007832323,0.000011249556,0.0010592259,0.00053195783,0.000047103753,0.00006237062,0.00015677372,0.00012833052,0.00017046489,0.012554524,0.9648488,0.020350857],"study_design_scores_gemma":[0.000020392787,0.0000024088606,0.001934596,0.0002610513,0.00001836071,0.000024475521,0.00012999155,0.000085755935,0.0001403983,0.0017520825,0.9956103,0.000020301159],"about_ca_topic_score_codex":0.9089114,"about_ca_topic_score_gemma":0.925224,"teacher_disagreement_score":0.9963596,"about_ca_system_score_codex":0.040556844,"about_ca_system_score_gemma":0.14636299,"threshold_uncertainty_score":0.6372819},"labels":[],"label_agreement":null},{"id":"W4409655196","doi":"10.2196/58567","title":"Patient-Related Metadata Reported in Sequencing Studies of SARS-CoV-2: Protocol for a Scoping Review and Bibliometric Analysis","year":2025,"lang":"en","type":"review","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"National Institute of Allergy and Infectious Diseases","keywords":"Metadata; Data extraction; MEDLINE; Data science; Computer science; GenBank; Information retrieval; World Wide Web; Biology; Genetics","score_opus":0.6113161080053996,"score_gpt":0.6692299291618927,"score_spread":0.05791382115649313,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409655196","genre_codex":"protocol","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":"protocol","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013563139,0.003159221,0.0032741185,0.0012671469,0.0003559454,0.9666254,0.021378877,0.00031846002,0.0022645916],"genre_scores_gemma":[0.0010655688,0.0017012618,0.007003662,0.0003197543,0.000053969186,0.9860241,0.003211863,0.000036168673,0.0005837498],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.9281333,0.018520288,0.037971355,0.0039786343,0.008688123,0.0027082292],"domain_scores_gemma":[0.8171573,0.058874637,0.036025915,0.015751526,0.06789638,0.004294178],"candidate_categories":["metaresearch","bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.10674366,0.00424369,0.012465139,0.04260089,0.0045099985,0.007327716,0.0044376776,0.0062231473,0.051781606],"category_scores_gemma":[0.15159114,0.0037094355,0.011857981,0.033385724,0.004136299,0.007623248,0.007471452,0.0038742167,0.010488561],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00517408,0.0005321148,0.0037017881,0.7885736,0.0020192803,0.0015633871,0.005117286,0.0016686134,0.0038023908,0.005225837,0.09096412,0.09165755],"study_design_scores_gemma":[0.007830044,0.0015665765,0.017776795,0.43539578,0.005363806,0.0009956204,0.0047331783,0.0013257707,0.0043705287,0.011030298,0.50890136,0.00071028015],"about_ca_topic_score_codex":0.005671734,"about_ca_topic_score_gemma":0.009968993,"teacher_disagreement_score":0.95739913,"about_ca_system_score_codex":0.012906103,"about_ca_system_score_gemma":0.0705534,"threshold_uncertainty_score":0.5645212},"labels":[],"label_agreement":null},{"id":"W4409733418","doi":"10.63485/v2qpb-3j174","title":"OA could solve problem for Canadian dissertations","year":2004,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Computer science; Political science","score_opus":0.023460832967553434,"score_gpt":0.30624119514644743,"score_spread":0.282780362178894,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409733418","genre_codex":"other","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.020122869,0.0048697414,0.03095861,0.17934069,0.008702801,0.00078129454,0.006061089,0.0038878333,0.7452752],"genre_scores_gemma":[0.16845956,0.0052212966,0.054638248,0.01877799,0.0016749221,0.0005210766,0.008833533,0.0023602017,0.7395131],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98044246,0.002828463,0.0013591384,0.0031093252,0.009405722,0.0028548092],"domain_scores_gemma":[0.9441664,0.0060673556,0.0017609992,0.0072086067,0.034704085,0.0060925665],"candidate_categories":["open_science"],"consensus_categories":[],"category_scores_codex":[0.01730062,0.00085932715,0.0011173083,0.0073580686,0.022728356,0.018135874,0.003639561,0.004023137,0.084652826],"category_scores_gemma":[0.07754845,0.00075561047,0.0018465399,0.014018607,0.0044082105,0.012274078,0.008897409,0.0048170188,0.021449694],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016798015,0.00007200307,0.0033246975,0.00030258248,0.000038836486,0.0005992056,0.006732966,0.0004773124,0.00071153394,0.37277538,0.4956758,0.11912174],"study_design_scores_gemma":[0.000022652503,0.000010989139,0.0014828041,0.0002420402,0.000033049026,0.00019788083,0.0037567983,0.000689302,0.00036312768,0.02512565,0.9680341,0.000041661388],"about_ca_topic_score_codex":0.63125485,"about_ca_topic_score_gemma":0.68425405,"teacher_disagreement_score":0.9963604,"about_ca_system_score_codex":0.044422343,"about_ca_system_score_gemma":0.092298254,"threshold_uncertainty_score":0.7418335},"labels":[],"label_agreement":null},{"id":"W4409818778","doi":"10.63485/35txp-yhe04","title":"More on OA to Canadian dissertations","year":2004,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Library science; Political science; Computer science","score_opus":0.017803906859458967,"score_gpt":0.31254109192332963,"score_spread":0.29473718506387064,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409818778","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008603627,0.004478535,0.0010242747,0.06348132,0.023326904,0.00025248463,0.016488355,0.0028072956,0.8872806],"genre_scores_gemma":[0.0026997144,0.0020585624,0.0010395807,0.008386176,0.0027032015,0.00008274637,0.004339331,0.0014454866,0.9772453],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99199986,0.00042774907,0.00023045906,0.000586619,0.005360358,0.0013948227],"domain_scores_gemma":[0.9667281,0.0020567158,0.0006240676,0.0022606857,0.019384952,0.008945598],"candidate_categories":["scholarly_communication","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0052674566,0.0012496897,0.0016257532,0.0105908355,0.017098976,0.015690919,0.0017833606,0.0030871183,0.64237964],"category_scores_gemma":[0.022672454,0.0007467748,0.0014046317,0.020274632,0.0028045774,0.006611308,0.00677659,0.0047651534,0.31608188],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000006869768,0.000008911921,0.00006487381,0.00003033128,0.0000011544707,0.000024363573,0.00014078685,0.000009821681,0.000046767596,0.0022300852,0.989755,0.0076811076],"study_design_scores_gemma":[0.0000028259722,0.0000017145652,0.0002846866,0.000039101407,9.3453303e-7,0.000007690937,0.00019007616,0.000005693164,0.000015378097,0.0003105269,0.9991335,0.000007831813],"about_ca_topic_score_codex":0.71232814,"about_ca_topic_score_gemma":0.88394237,"teacher_disagreement_score":0.9982166,"about_ca_system_score_codex":0.037070375,"about_ca_system_score_gemma":0.056338295,"threshold_uncertainty_score":0.578732},"labels":[],"label_agreement":null},{"id":"W4409948370","doi":"10.1093/nar/gkaf337","title":"OntoTiger: a platform of ontology-based application tools for integrative biomedical exploration","year":2025,"lang":"en","type":"article","venue":"Nucleic Acids Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Science North","funders":"National Natural Science Foundation of China","keywords":"Ontology; Open Biomedical Ontologies; Ontology-based data integration; Pairwise comparison; Computer science; Gene ontology; Process ontology; Upper ontology; Annotation; Ontology alignment; Suggested Upper Merged Ontology; Information retrieval; Similarity (geometry); Semantic Web; Bioinformatics; Biology; Artificial intelligence; Gene","score_opus":0.08527833998254629,"score_gpt":0.4058098358403694,"score_spread":0.3205314958578231,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4409948370","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006337192,0.0009149071,0.6320432,0.00067674404,0.00022616747,0.00096452615,0.034400705,0.31043774,0.01399884],"genre_scores_gemma":[0.06959695,0.0033663104,0.63929045,0.001637505,0.00024801574,0.003505512,0.21824707,0.046414386,0.017693823],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99871886,0.00016943457,0.00018027496,0.00033314567,0.0004728042,0.00012551108],"domain_scores_gemma":[0.9983571,0.00063189055,0.0001990942,0.00033883314,0.00024714947,0.0002258446],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0024342465,0.002455948,0.0015180266,0.0049499725,0.0009986522,0.0034464265,0.0030147128,0.0012448574,0.016767856],"category_scores_gemma":[0.0052309674,0.0011556569,0.0026668643,0.003501468,0.00090644695,0.004373304,0.0066118753,0.0024984337,0.010775795],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021016416,0.00064968714,0.0057078493,0.006408499,0.0014505876,0.0029495712,0.0028005552,0.015587289,0.08130345,0.07250172,0.37371814,0.43482107],"study_design_scores_gemma":[0.0006577555,0.00024246165,0.008280189,0.0007219429,0.00039486727,0.0017033272,0.0005020016,0.111917004,0.033156045,0.099996395,0.741943,0.0004850116],"about_ca_topic_score_codex":0.0032618525,"about_ca_topic_score_gemma":0.0031083787,"teacher_disagreement_score":0.016767856,"about_ca_system_score_codex":0.001164067,"about_ca_system_score_gemma":0.0028173968,"threshold_uncertainty_score":0.05609405},"labels":[],"label_agreement":null},{"id":"W4410074846","doi":"10.1101/2025.05.02.25326887","title":"Evaluating the Potential of AI-Generated Synthetic Diaries in Parkinson’s Disease Research","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"Canadian Institutes of Health Research","keywords":"Parkinson's disease; Disease; Psychology; Data science; Medicine; Computer science; Internal medicine","score_opus":0.0853112994754933,"score_gpt":0.41090598203513135,"score_spread":0.325594682559638,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410074846","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7547957,0.0010802903,0.20116049,0.0044121784,0.0004979886,0.003456523,0.011174337,0.0037585634,0.019663818],"genre_scores_gemma":[0.8389282,0.00032870585,0.15131353,0.00051748654,0.000059697028,0.0015377584,0.006067906,0.00014159565,0.001105062],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9882983,0.009438607,0.0005759826,0.00070231035,0.0008491191,0.00013570955],"domain_scores_gemma":[0.89177877,0.09290657,0.0027059945,0.0079313945,0.003655033,0.0010222951],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01849574,0.00057565246,0.0003636031,0.0012368053,0.00047469317,0.0027798156,0.0014881081,0.0010799746,0.0027181169],"category_scores_gemma":[0.105274975,0.00023815947,0.00059585524,0.0010456229,0.0009315858,0.0019485246,0.002496393,0.00097093673,0.0008947414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005037981,0.0027867479,0.12207726,0.0056908447,0.00066650746,0.0017965009,0.024126988,0.1918713,0.013673017,0.043810237,0.024241064,0.56422156],"study_design_scores_gemma":[0.0011797764,0.0038945638,0.038068343,0.0018085489,0.0003568492,0.0013256122,0.008845939,0.77375734,0.017444966,0.06507511,0.087836064,0.00040685458],"about_ca_topic_score_codex":0.0019423548,"about_ca_topic_score_gemma":0.0023989666,"teacher_disagreement_score":0.01849574,"about_ca_system_score_codex":0.0011365089,"about_ca_system_score_gemma":0.0012081929,"threshold_uncertainty_score":0.09781599},"labels":[],"label_agreement":null},{"id":"W4410274719","doi":"10.1101/2025.05.09.25327342","title":"Linking international registries to FHIR and Phenopackets with RareLink: a scalable REDCap-based framework for rare disease data interoperability","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Dalhousie University; Alberta Children's Hospital; University of Calgary","funders":"","keywords":"Interoperability; Scalability; Computer science; Rare disease; Database; World Wide Web; Disease; Medicine; Internal medicine","score_opus":0.04669301839104795,"score_gpt":0.3339551727122725,"score_spread":0.28726215432122454,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410274719","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0045918524,0.00053475157,0.86609524,0.0017494182,0.0002823064,0.0016508809,0.015766839,0.10194939,0.0073793707],"genre_scores_gemma":[0.051664066,0.0011377238,0.82266855,0.0023390616,0.00019605858,0.0024825016,0.100182734,0.013378388,0.0059509138],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.98962164,0.0034034369,0.0018333744,0.0022516872,0.0022015036,0.0006883033],"domain_scores_gemma":[0.9816482,0.0070519797,0.0016930449,0.0059755184,0.0023279635,0.0013033503],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024002483,0.0018094315,0.0013324958,0.0075340765,0.0016476201,0.009285697,0.0055018826,0.002454243,0.006881232],"category_scores_gemma":[0.034735393,0.0016461452,0.0043456075,0.004840827,0.0019781026,0.009881059,0.0146937715,0.0037263439,0.0056124204],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016106877,0.0009705203,0.02499431,0.0044503063,0.0014219246,0.003601875,0.008126705,0.04587281,0.013829823,0.3071567,0.2726177,0.31534657],"study_design_scores_gemma":[0.00040054836,0.0002672485,0.007562342,0.0019123595,0.00040522296,0.0016487442,0.0012275929,0.1169665,0.015152514,0.15828773,0.69550174,0.00066738395],"about_ca_topic_score_codex":0.019180734,"about_ca_topic_score_gemma":0.015637226,"teacher_disagreement_score":0.024002483,"about_ca_system_score_codex":0.0029323017,"about_ca_system_score_gemma":0.008822843,"threshold_uncertainty_score":0.12693876},"labels":[],"label_agreement":null},{"id":"W4410347660","doi":"10.1101/2025.05.13.653828","title":"Large Language Models Can Extract Metadata for Annotation of Human Neuroimaging Publications","year":2025,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Institute on Drug Abuse","keywords":"Metadata; Annotation; Neuroimaging; Computer science; Natural language processing; Information retrieval; Data science; Artificial intelligence; World Wide Web; Psychology; Neuroscience","score_opus":0.030103732778211843,"score_gpt":0.2940016720894056,"score_spread":0.26389793931119376,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410347660","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.047604494,0.0025596288,0.784364,0.004715575,0.0008869548,0.00074070075,0.028360521,0.11449309,0.016275084],"genre_scores_gemma":[0.26567495,0.0010979106,0.6517713,0.0015128532,0.00035505244,0.001117773,0.06465247,0.0063340836,0.007483598],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9898453,0.0052814046,0.0009911492,0.0016172668,0.0019648003,0.0003001233],"domain_scores_gemma":[0.9472686,0.033944663,0.0029589918,0.009347568,0.005742274,0.0007378759],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0153002385,0.002090791,0.0006857021,0.0062356894,0.0011209068,0.004615332,0.0019659428,0.0020506557,0.0065781544],"category_scores_gemma":[0.072028115,0.00088380644,0.0017991071,0.0036043052,0.00091872213,0.006830748,0.004426381,0.0024427578,0.011303342],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014559194,0.0005667753,0.02242414,0.0039422866,0.00090940966,0.00072692876,0.00309894,0.037364513,0.030811675,0.026105035,0.1892518,0.6833426],"study_design_scores_gemma":[0.00021308691,0.00038789152,0.014171107,0.0009012769,0.000436521,0.000719277,0.0016307764,0.6073826,0.05238222,0.107225515,0.21416031,0.00038942418],"about_ca_topic_score_codex":0.005057836,"about_ca_topic_score_gemma":0.017779283,"teacher_disagreement_score":0.9846998,"about_ca_system_score_codex":0.0017970898,"about_ca_system_score_gemma":0.0039978772,"threshold_uncertainty_score":0.080916405},"labels":[],"label_agreement":null},{"id":"W4410397499","doi":"10.1016/j.sapharm.2025.05.008","title":"Consistency of Medical Subject Headings assignment: A test-retest reliability analysis","year":2025,"lang":"en","type":"article","venue":"Research in Social and Administrative Pharmacy","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Subject (documents); Reliability (semiconductor); Test (biology); Information retrieval; Computer science; Psychology; Reliability engineering; Artificial intelligence; World Wide Web; Engineering; Biology","score_opus":0.17539365046157487,"score_gpt":0.5228577643158194,"score_spread":0.3474641138542446,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410397499","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7985722,0.01192886,0.14187783,0.0012542203,0.0024428987,0.021854691,0.0059086108,0.0013397328,0.0148209715],"genre_scores_gemma":[0.89773977,0.001237869,0.079001054,0.00034550394,0.0003568075,0.016298607,0.002882028,0.00039692866,0.0017413612],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.814525,0.08559629,0.043171924,0.017882144,0.036981847,0.0018427415],"domain_scores_gemma":[0.53743255,0.27688628,0.043602187,0.030913692,0.1096916,0.0014737775],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.15739448,0.0010710045,0.0030346583,0.009546117,0.0016672355,0.0033175794,0.0016159058,0.0012059147,0.0018354197],"category_scores_gemma":[0.35331416,0.0008955316,0.0051865405,0.0069015664,0.0028123895,0.0033164488,0.004376858,0.0013266609,0.0009663721],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.009013396,0.001040454,0.67924565,0.015170594,0.016346442,0.00042551098,0.03382018,0.00343681,0.01386855,0.0043085148,0.011056105,0.21226774],"study_design_scores_gemma":[0.0018771962,0.0056991735,0.8664946,0.0049311207,0.010382059,0.001275712,0.009154501,0.022761663,0.022181608,0.011696014,0.042744722,0.0008016068],"about_ca_topic_score_codex":0.0015830669,"about_ca_topic_score_gemma":0.0026919253,"teacher_disagreement_score":0.84260553,"about_ca_system_score_codex":0.0017328077,"about_ca_system_score_gemma":0.0025788024,"threshold_uncertainty_score":0.83239156},"labels":[],"label_agreement":null},{"id":"W4410498594","doi":"10.1111/bcpt.70054","title":"Deprescribed or Discontinued? Addressing Terminology Misinterpretation in Research","year":2025,"lang":"en","type":"article","venue":"Basic & Clinical Pharmacology & Toxicology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université Laval; Quebec Network for Research on Aging","funders":"","keywords":"Terminology; Medicine; Psychology; Philosophy; Linguistics","score_opus":0.2829110607444573,"score_gpt":0.5578423064157094,"score_spread":0.2749312456712521,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410498594","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15103991,0.05679864,0.3517919,0.34732842,0.016122328,0.0011293172,0.008924131,0.0017882014,0.065077215],"genre_scores_gemma":[0.68887824,0.021897696,0.22430934,0.05124649,0.0036994051,0.0008141027,0.0041422443,0.0010750195,0.003937534],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8907374,0.0753043,0.016184289,0.0037238426,0.012894213,0.0011558522],"domain_scores_gemma":[0.5825821,0.33506125,0.03315697,0.018575842,0.028660243,0.0019636052],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.106221095,0.00072353316,0.00094880926,0.0077076657,0.0019423527,0.006363003,0.0024441262,0.0023413973,0.003631153],"category_scores_gemma":[0.3156235,0.0005517858,0.0012407363,0.006796672,0.006591342,0.011555677,0.0078091766,0.0037427135,0.0012535296],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009213476,0.00010828412,0.04564334,0.009520179,0.00048582593,0.009561528,0.11814274,0.0012470423,0.006907627,0.3026295,0.11294578,0.3918868],"study_design_scores_gemma":[0.00008739239,0.0001144633,0.008457842,0.013681974,0.0007040647,0.011476498,0.046871617,0.0040940503,0.010431882,0.26802936,0.6358109,0.00024000535],"about_ca_topic_score_codex":0.0043468704,"about_ca_topic_score_gemma":0.0055676857,"teacher_disagreement_score":0.8937789,"about_ca_system_score_codex":0.004637204,"about_ca_system_score_gemma":0.010664971,"threshold_uncertainty_score":0.56175756},"labels":[],"label_agreement":null},{"id":"W4410539343","doi":"10.1093/database/baaf008","title":"An exploratory study combining Virtual Reality and Semantic Web for life science research using Graph2VR","year":2025,"lang":"en","type":"article","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"H2020 Health","keywords":"Computer science; Virtual reality; World Wide Web; Exploratory research; Semantic Web; Information retrieval; Web of science; Social Semantic Web; Data science; Human–computer interaction; MEDLINE; Sociology","score_opus":0.14193607316905618,"score_gpt":0.44667886393519246,"score_spread":0.3047427907661363,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410539343","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7209753,0.0015768682,0.2023653,0.0037661847,0.00023251133,0.003313334,0.02463509,0.0123540135,0.030781444],"genre_scores_gemma":[0.56185704,0.0008555563,0.4140719,0.000811745,0.00007263291,0.0021476182,0.013983979,0.002093436,0.004106104],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99526346,0.0033759912,0.0002002752,0.00044920007,0.00052650826,0.00018455759],"domain_scores_gemma":[0.9678908,0.026606888,0.0005083749,0.0029575804,0.0013523594,0.0006840153],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0096187005,0.00069506286,0.00048486388,0.0036248565,0.0010168629,0.0036310875,0.0016515715,0.0011148464,0.0062185926],"category_scores_gemma":[0.022461317,0.00041574694,0.0012175953,0.0036874902,0.0010212983,0.0038539208,0.00362178,0.0011054665,0.0012591154],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039159455,0.0053337445,0.11262183,0.010624775,0.0009934258,0.0097982,0.094558455,0.034858964,0.043057144,0.060977,0.07321539,0.55004513],"study_design_scores_gemma":[0.0014091077,0.004583269,0.081579834,0.001873649,0.00089583796,0.005945739,0.059940048,0.15793121,0.04314957,0.079428986,0.5626083,0.00065452274],"about_ca_topic_score_codex":0.0045828917,"about_ca_topic_score_gemma":0.009628153,"teacher_disagreement_score":0.0096187005,"about_ca_system_score_codex":0.0010485747,"about_ca_system_score_gemma":0.0011344213,"threshold_uncertainty_score":0.050869167},"labels":[],"label_agreement":null},{"id":"W4410673223","doi":"10.1075/term.00083.vid","title":"AI as a resource for the clarification of medical terminology","year":2025,"lang":"en","type":"article","venue":"Terminology International Journal of Theoretical and Applied Issues in Specialized Communication","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Terminology; Resource (disambiguation); Computer science; Management science; Philosophy; Linguistics; Engineering","score_opus":0.013415646633596866,"score_gpt":0.356907057420703,"score_spread":0.34349141078710616,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410673223","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.054387704,0.0044270633,0.8040373,0.012981669,0.0012830484,0.0015696671,0.00611905,0.0036174739,0.11157706],"genre_scores_gemma":[0.24992506,0.0021619315,0.7257179,0.0012114194,0.00030537997,0.0013201911,0.0066313203,0.00079022977,0.011936559],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9913565,0.0055110743,0.0013132029,0.00047579315,0.0012232158,0.000120199176],"domain_scores_gemma":[0.9581059,0.030686282,0.0020673454,0.0051810346,0.0035098682,0.00044968937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008248996,0.00070143875,0.00066798675,0.009152101,0.001988113,0.005498795,0.001507117,0.0009647442,0.0109066935],"category_scores_gemma":[0.028772322,0.00046531463,0.00049891806,0.005907674,0.0034438712,0.010875012,0.0060160295,0.0027674853,0.0026701016],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023880892,0.00010656531,0.001862168,0.0036964372,0.000055041935,0.0014055015,0.05144608,0.0020803446,0.027445773,0.5857559,0.048364118,0.27754325],"study_design_scores_gemma":[0.000045964407,0.0001309152,0.004437454,0.0036500355,0.00011436728,0.0023506074,0.025568694,0.018498322,0.0142724365,0.17461462,0.7561366,0.0001799569],"about_ca_topic_score_codex":0.0018091542,"about_ca_topic_score_gemma":0.0020135457,"teacher_disagreement_score":0.0109066935,"about_ca_system_score_codex":0.0018926798,"about_ca_system_score_gemma":0.0024794056,"threshold_uncertainty_score":0.043625414},"labels":[],"label_agreement":null},{"id":"W4410811567","doi":"10.1200/jco.2025.43.16_suppl.e23110","title":"Assessment of OncoQuébec, an oncology clinical trials search engine, on clinical trial recruitment in Quebec.","year":2025,"lang":"en","type":"article","venue":"Journal of Clinical Oncology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Jewish General Hospital; Quebec - Clinical Research Organization in Cancer","funders":"Innovative Medicines Canada; AbbVie; Pfizer","keywords":"Medicine; Clinical trial; Clinical Oncology; Oncology; Internal medicine; Family medicine; Cancer","score_opus":0.5833898828062747,"score_gpt":0.6698820581182794,"score_spread":0.0864921753120047,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410811567","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.38500673,0.09122978,0.006123855,0.06396345,0.0009923554,0.012026612,0.37643543,0.0037588682,0.060462896],"genre_scores_gemma":[0.8719666,0.011890339,0.018781863,0.013568325,0.00030190372,0.007859041,0.06839204,0.0007280089,0.0065119388],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9011903,0.060155664,0.0126188435,0.003887837,0.01922732,0.0029200637],"domain_scores_gemma":[0.28857166,0.5562544,0.053558417,0.010435859,0.08049504,0.010684545],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14826673,0.00088474166,0.001828747,0.009468904,0.0023432416,0.007154733,0.003395565,0.0019973693,0.012920011],"category_scores_gemma":[0.5658742,0.00068760896,0.0027236561,0.018648896,0.0012342811,0.005505526,0.0051781526,0.0015873912,0.0014936542],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.01261055,0.00052330823,0.4352209,0.040770322,0.005514248,0.0004525017,0.005411934,0.002953469,0.0006897495,0.0029394852,0.2714821,0.22143133],"study_design_scores_gemma":[0.0034399687,0.0020904338,0.8069644,0.019242885,0.0055104285,0.0003907514,0.004968721,0.015279188,0.00071683957,0.0011637352,0.13970989,0.0005228198],"about_ca_topic_score_codex":0.7476573,"about_ca_topic_score_gemma":0.8771985,"teacher_disagreement_score":0.9600223,"about_ca_system_score_codex":0.03997773,"about_ca_system_score_gemma":0.08124666,"threshold_uncertainty_score":0.7841188},"labels":[],"label_agreement":null},{"id":"W4410872475","doi":"10.1101/2025.05.28.25328511","title":"PheCode-guided multi-modal topic modeling of electronic health records improves disease incidence prediction and GWAS discovery from UK Biobank","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal; McGill University","funders":"Fonds de recherche du Québec – Nature et technologies; Alliance de recherche numérique du Canada; Canada First Research Excellence Fund; Natural Sciences and Engineering Research Council of Canada; McGill University","keywords":"Biobank; Health records; Modal; Data science; Disease; Incidence (geometry); Genome-wide association study; Data discovery; Medicine; Computer science; Internal medicine; Bioinformatics; Political science; World Wide Web; Health care; Biology; Mathematics; Genetics","score_opus":0.026758487194818263,"score_gpt":0.30817207759478377,"score_spread":0.2814135903999655,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410872475","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.18604372,0.0016369502,0.75360715,0.0043008043,0.00021751074,0.00033746602,0.0349777,0.013778182,0.005100478],"genre_scores_gemma":[0.63506323,0.0009773499,0.28692046,0.0013004963,0.00038503908,0.00071325473,0.069088995,0.0013313983,0.004219803],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9974342,0.0013625645,0.00015868756,0.0006652909,0.00024732912,0.00013193277],"domain_scores_gemma":[0.9873111,0.010231447,0.00062725437,0.0010669179,0.0005876846,0.0001756038],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0049134097,0.0007197127,0.0009346401,0.0018092748,0.0005467284,0.0017658646,0.00092319463,0.0011644263,0.0034081677],"category_scores_gemma":[0.024459358,0.00048624125,0.0019749233,0.0017221917,0.0003745285,0.0015397151,0.0019754327,0.001580873,0.0016934219],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0024085618,0.0005421389,0.1891331,0.001577614,0.0020213411,0.0008103354,0.002547827,0.2131647,0.010918516,0.02976461,0.0968636,0.4502476],"study_design_scores_gemma":[0.00027387287,0.00011258429,0.022437412,0.00016037923,0.00026756714,0.00029442238,0.00021878949,0.9152316,0.0034072264,0.03785674,0.019645358,0.00009406833],"about_ca_topic_score_codex":0.011411683,"about_ca_topic_score_gemma":0.023268476,"teacher_disagreement_score":0.011411683,"about_ca_system_score_codex":0.0007060873,"about_ca_system_score_gemma":0.0014710117,"threshold_uncertainty_score":0.025984883},"labels":[],"label_agreement":null},{"id":"W4410984146","doi":"10.22541/au.174897564.48312932/v1","title":"The STARR Protocol: An Automated LLM Methodology for Enhanced Systematic Literature Review Screening","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College","funders":"","keywords":"Protocol (science); Systematic review; Computer science; Medicine; Political science; MEDLINE; Alternative medicine; Law","score_opus":0.07456971998791556,"score_gpt":0.440407428649505,"score_spread":0.36583770866158943,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4410984146","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034582745,0.0025663257,0.5704427,0.0046239574,0.0019411312,0.14814186,0.18513809,0.07634482,0.0073429346],"genre_scores_gemma":[0.0042331875,0.0005053136,0.7727998,0.00050345453,0.00017362373,0.19975512,0.017630735,0.0020739369,0.002324822],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.89880854,0.055989802,0.032949913,0.0062101884,0.005083942,0.0009575414],"domain_scores_gemma":[0.65629184,0.24805388,0.022677045,0.035441514,0.034518078,0.0030176553],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.14460672,0.004065623,0.0045347195,0.018168483,0.003026098,0.007375475,0.0041772127,0.003637953,0.1614547],"category_scores_gemma":[0.30991074,0.0046269065,0.008537398,0.012524304,0.0026815608,0.004843952,0.010696083,0.005034384,0.037364602],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0039744494,0.0003738005,0.0032987886,0.18126097,0.0040720096,0.00069469254,0.0047969683,0.003925118,0.008939686,0.035544038,0.41685018,0.33626938],"study_design_scores_gemma":[0.011490139,0.0008159649,0.006297539,0.034448534,0.0062595834,0.00077128457,0.0015308567,0.026951162,0.0152074285,0.10037161,0.7944803,0.001375528],"about_ca_topic_score_codex":0.0028076796,"about_ca_topic_score_gemma":0.007965868,"teacher_disagreement_score":0.8553933,"about_ca_system_score_codex":0.003557952,"about_ca_system_score_gemma":0.04090593,"threshold_uncertainty_score":0.7647626},"labels":[],"label_agreement":null},{"id":"W4411019801","doi":"10.1109/jtehm.2025.3576570","title":"Improving Transformer Performance for French Clinical Notes Classification Using Mixture of Experts on a Limited Dataset","year":2025,"lang":"en","type":"article","venue":"IEEE Journal of Translational Engineering in Health and Medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Centre Hospitalier Universitaire Sainte-Justine; Université de Montréal; École de Technologie Supérieure","funders":"Fonds de recherche du Québec – Nature et technologies; Natural Sciences and Engineering Research Council of Canada; Fonds de Recherche du Québec - Santé; Institut de Valorisation des Données","keywords":"Computer science; Transformer; Artificial intelligence; Natural language processing; Pattern recognition (psychology); Machine learning; Engineering","score_opus":0.0660339813952497,"score_gpt":0.382899591068407,"score_spread":0.3168656096731573,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411019801","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.53484833,0.0092449365,0.40605032,0.004488842,0.0009398443,0.000704523,0.009302689,0.020043297,0.014377205],"genre_scores_gemma":[0.86219907,0.000981876,0.11205657,0.0012326683,0.00023411149,0.00018872482,0.0152882775,0.0003003727,0.0075183576],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9984079,0.00063401775,0.00012394872,0.00047398548,0.00020211673,0.00015812456],"domain_scores_gemma":[0.9964192,0.0023501532,0.00011948559,0.00039323556,0.00057572173,0.00014217451],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0039803907,0.001819626,0.0011004778,0.0022490078,0.0005939289,0.0013978686,0.0013819747,0.0017680312,0.0027054907],"category_scores_gemma":[0.009433242,0.0003302695,0.0016335252,0.0011291298,0.00043340915,0.0018390119,0.0013770654,0.0019247092,0.0023340264],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022063064,0.0007463618,0.030602008,0.00047840414,0.00063067913,0.00093607686,0.00045655348,0.23108128,0.007557447,0.003477382,0.04723278,0.6745947],"study_design_scores_gemma":[0.00007025331,0.00018833079,0.0025359192,0.00004382939,0.000085762564,0.00023286662,0.00017760404,0.98564065,0.0040637166,0.002855488,0.0040639895,0.00004153413],"about_ca_topic_score_codex":0.022026712,"about_ca_topic_score_gemma":0.028405793,"teacher_disagreement_score":0.022026712,"about_ca_system_score_codex":0.001330272,"about_ca_system_score_gemma":0.0020213018,"threshold_uncertainty_score":0.043797016},"labels":[],"label_agreement":null},{"id":"W4411189102","doi":"10.1186/s40708-025-00261-2","title":"Detecting label noise in longitudinal Alzheimer’s data with explainable artificial intelligence","year":2025,"lang":"en","type":"article","venue":"Brain Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Canadian Institutes of Health Research; Genentech; National Institutes of Health; Northern California Institute for Research and Education; Servier; BioClinica; U.S. Department of Defense; Alzheimer's Disease Neuroimaging Initiative; University of Southern California; Bristol-Myers Squibb; Eli Lilly and Company; Biogen; Eisai; National Institute on Aging; Alzheimer's Association","keywords":"Noise (video); Artificial intelligence; Pattern recognition (psychology); Longitudinal data; Computer science; Machine learning; Data mining","score_opus":0.08247190702132093,"score_gpt":0.3400355710690964,"score_spread":0.2575636640477754,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411189102","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4745375,0.00040324518,0.5223004,0.0011481039,0.000037470105,0.000082575505,0.00041808837,0.00047144588,0.00060117344],"genre_scores_gemma":[0.94561774,0.000071337716,0.053301033,0.00014354201,0.0000324437,0.00006679456,0.0005732387,0.000024552248,0.00016921708],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954809,0.0026886028,0.00032796664,0.0008748012,0.00047762436,0.00015007347],"domain_scores_gemma":[0.9457026,0.043015666,0.0048285755,0.004959819,0.0012425662,0.00025075607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013555287,0.0006435301,0.0007827157,0.002014325,0.0007622507,0.0020504692,0.0010527375,0.0011909151,0.00042669],"category_scores_gemma":[0.053738434,0.00037389898,0.0008953626,0.0012234779,0.0016446655,0.0018210921,0.0018116385,0.0018003029,0.000101232785],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008540223,0.00046207692,0.21499515,0.00028206196,0.0008167024,0.0006048369,0.0027283218,0.5608344,0.0070425305,0.02367973,0.0015877875,0.18611237],"study_design_scores_gemma":[0.000017936225,0.00011026613,0.021239806,0.00004095612,0.00006493138,0.00007878773,0.00015022117,0.92823845,0.0021035199,0.04731227,0.00061182864,0.00003112582],"about_ca_topic_score_codex":0.0036744045,"about_ca_topic_score_gemma":0.0045688674,"teacher_disagreement_score":0.013555287,"about_ca_system_score_codex":0.0011820468,"about_ca_system_score_gemma":0.00092693005,"threshold_uncertainty_score":0.071688056},"labels":[],"label_agreement":null},{"id":"W4411272807","doi":"10.1073/pnas.2501660122","title":"Take caution in using LLMs as human surrogates","year":2025,"lang":"en","type":"article","venue":"Proceedings of the National Academy of Sciences","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":29,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Quest University Canada","funders":"","keywords":"Psychology; Political science","score_opus":0.051398931778552225,"score_gpt":0.3707658822394894,"score_spread":0.3193669504609372,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411272807","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039884705,0.005427798,0.60821575,0.25034842,0.010193538,0.0011822563,0.0026647663,0.004484504,0.0775982],"genre_scores_gemma":[0.4742591,0.0016848261,0.42896077,0.07056944,0.0018511122,0.0040940316,0.0013031457,0.0018621074,0.01541538],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.84103185,0.1331422,0.0054955687,0.0068423403,0.0126713095,0.00081677013],"domain_scores_gemma":[0.55533016,0.32228473,0.016552076,0.07322832,0.028841153,0.003763544],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.13397126,0.0018145846,0.001870335,0.0024919112,0.002746491,0.010554662,0.007367279,0.0051721334,0.0090487115],"category_scores_gemma":[0.5460492,0.0013183628,0.0015320289,0.0021336838,0.012478318,0.015142387,0.008034187,0.011985687,0.005976952],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0018976729,0.00032369027,0.016473653,0.0022625912,0.0009608275,0.00085821946,0.022677954,0.015547455,0.002447133,0.62442553,0.16417827,0.14794701],"study_design_scores_gemma":[0.00023952436,0.00026490944,0.0028812897,0.0018528209,0.0001759431,0.00057952054,0.0044983486,0.033305224,0.0037954273,0.73815703,0.21400253,0.00024743003],"about_ca_topic_score_codex":0.0055879215,"about_ca_topic_score_gemma":0.0070115533,"teacher_disagreement_score":0.13397126,"about_ca_system_score_codex":0.0025396782,"about_ca_system_score_gemma":0.003727769,"threshold_uncertainty_score":0.70851624},"labels":[],"label_agreement":null},{"id":"W4411379405","doi":"10.1007/978-981-96-8186-0_17","title":"Improving Clinical Note Generation from Complex Doctor-Patient Conversation","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Tellabs (Canada); Université de Montréal; Mila - Quebec Artificial Intelligence Institute; Université du Québec à Montréal","funders":"","keywords":"Conversation; Computer science; Programming language; Artificial intelligence; Natural language processing; Linguistics; Philosophy","score_opus":0.036736752718166545,"score_gpt":0.30887766094679525,"score_spread":0.2721409082286287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411379405","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10008024,0.0012727516,0.8120532,0.001888459,0.0012574008,0.0014288863,0.012170265,0.060889363,0.008959476],"genre_scores_gemma":[0.31151548,0.0005271045,0.6469239,0.0007070264,0.0004680041,0.0006297038,0.025452994,0.0024207528,0.011355014],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983966,0.0006104673,0.00011917147,0.00036325026,0.0004210591,0.000089486966],"domain_scores_gemma":[0.9860789,0.01134453,0.00025756622,0.00060759846,0.0014248347,0.0002865937],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020203297,0.0014874834,0.0010024494,0.0017085538,0.0006752885,0.0018252392,0.0015563768,0.0017437447,0.018340416],"category_scores_gemma":[0.013415979,0.0005531304,0.0008465044,0.0010591127,0.00024473818,0.0014619116,0.002066404,0.0014498092,0.011220268],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025565063,0.0004618354,0.0055075358,0.0006081105,0.00016856007,0.0016386141,0.0009205459,0.014284833,0.035987873,0.0023394516,0.065804295,0.86972183],"study_design_scores_gemma":[0.00047380026,0.0006115791,0.004108312,0.00015991466,0.00027081263,0.0017027057,0.0011181606,0.8740894,0.069075525,0.011595555,0.036663312,0.00013105676],"about_ca_topic_score_codex":0.002941674,"about_ca_topic_score_gemma":0.0031826787,"teacher_disagreement_score":0.018340416,"about_ca_system_score_codex":0.0005501033,"about_ca_system_score_gemma":0.0011322098,"threshold_uncertainty_score":0.061354756},"labels":[],"label_agreement":null},{"id":"W4411490737","doi":"10.1038/s41598-025-06447-2","title":"Evaluating language model embeddings for Parkinson’s disease cohort harmonization using a novel manually curated variable mapping schema","year":2025,"lang":"en","type":"article","venue":"Scientific Reports","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institute on Aging; National Institute of Biomedical Imaging and Bioengineering; Canadian Institutes of Health Research; Genentech; National Institutes of Health; IXICO; H. Lundbeck A/S; Mitsubishi Tanabe Pharma Corporation; Servier; Université de Genève; Shionogi; Japan Science and Technology Agency; Astellas Pharma; Fondazione Cariplo; Eisai; Daiichi-Sankyo; European Commission; GHR Foundation; Pfizer; Biogen; BioClinica; F. Hoffmann-La Roche; Wellcome Trust; University of Southern California; U.S. Department of Defense; Eli Lilly and Company; Bristol-Myers Squibb; Meso Scale Diagnostics; Alzheimer's Disease Neuroimaging Initiative; Ministero della Salute; Novartis Pharmaceuticals Corporation; Alzheimer's Association; European Federation of Pharmaceutical Industries and Associations; Brigham and Women's Hospital","keywords":"Harmonization; Computer science; Cohort; Artificial intelligence; Schema (genetic algorithms); Language model; Natural language processing; Data mining; Machine learning; Medicine; Pathology","score_opus":0.04914273914335597,"score_gpt":0.3592620222019989,"score_spread":0.3101192830586429,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411490737","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6031591,0.0011889043,0.35607088,0.000855284,0.00029427838,0.0007104707,0.019147197,0.0147638535,0.003809978],"genre_scores_gemma":[0.61698896,0.00031668812,0.3395544,0.0002058478,0.000039721566,0.00041754343,0.04045587,0.0005674351,0.001453572],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9977719,0.00086275494,0.0003363174,0.0006387368,0.00030777615,0.00008232579],"domain_scores_gemma":[0.99447685,0.0032356272,0.00046744285,0.000920009,0.00078843563,0.00011164217],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042374264,0.00077143585,0.000349203,0.0023625877,0.0003667771,0.0012248649,0.0006862643,0.00076126127,0.0017589444],"category_scores_gemma":[0.016422862,0.00019558535,0.0010307025,0.0013443014,0.0003294475,0.0018660822,0.0014839618,0.0007862479,0.0008224255],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013299864,0.0008743593,0.10889871,0.0016458761,0.0009305229,0.0010944614,0.0024772931,0.15018214,0.032571774,0.01281016,0.030775053,0.6564096],"study_design_scores_gemma":[0.00017354042,0.00053234014,0.019830478,0.0002339004,0.0002524052,0.0008806541,0.0018514916,0.8973254,0.04108328,0.011755919,0.025990942,0.00008971797],"about_ca_topic_score_codex":0.0037925995,"about_ca_topic_score_gemma":0.004895324,"teacher_disagreement_score":0.0042374264,"about_ca_system_score_codex":0.0006558057,"about_ca_system_score_gemma":0.0012525533,"threshold_uncertainty_score":0.022409916},"labels":[],"label_agreement":null},{"id":"W4411507092","doi":"10.1007/978-3-031-95841-0_78","title":"Using Word Embeddings to Extract Semantic Relations from Biomedical Texts: Towards Literature-Based Discovery","year":2025,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; University of Ottawa","funders":"","keywords":"Computer science; Natural language processing; Word (group theory); Artificial intelligence; Information retrieval; Linguistics; Philosophy","score_opus":0.018760445088935696,"score_gpt":0.2938758378165231,"score_spread":0.2751153927275874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411507092","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.061255544,0.016800404,0.87520254,0.004037248,0.0014270265,0.0006908162,0.021215657,0.005973441,0.013397379],"genre_scores_gemma":[0.12488781,0.008951388,0.8271745,0.0005941773,0.00043052708,0.0005142328,0.032517456,0.0005002339,0.0044297203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9980573,0.00045554823,0.00043967148,0.0004692542,0.00050247746,0.000075852484],"domain_scores_gemma":[0.99542326,0.0025979474,0.00052912196,0.00041473153,0.00088953396,0.00014543305],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0017166964,0.0013124745,0.0010486735,0.012066142,0.0008080255,0.0036325557,0.0010293723,0.0012489561,0.0028470194],"category_scores_gemma":[0.009175564,0.00055705354,0.0017326361,0.012337071,0.00081110926,0.00640981,0.0030345672,0.0014666065,0.0041080797],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023850662,0.00040382237,0.0044819247,0.003123183,0.00038556618,0.0006996028,0.0015894735,0.0033925232,0.018439416,0.028290829,0.023565024,0.9153902],"study_design_scores_gemma":[0.0001827395,0.0004972119,0.012221025,0.0030584158,0.0015447449,0.003403122,0.005334867,0.22165889,0.036223218,0.39110765,0.32447377,0.00029443606],"about_ca_topic_score_codex":0.0013401812,"about_ca_topic_score_gemma":0.0026011458,"teacher_disagreement_score":0.012066142,"about_ca_system_score_codex":0.000629245,"about_ca_system_score_gemma":0.0020678914,"threshold_uncertainty_score":0.009524226},"labels":[],"label_agreement":null},{"id":"W4411622822","doi":"10.1093/gigascience/giaf070","title":"Extraction of biological terms using large language models enhances the usability of metadata in the BioSample database","year":2025,"lang":"en","type":"article","venue":"GigaScience","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute of Genetics; National Bioscience Database Center; Japan Society for the Promotion of Science; Japan Science and Technology Corporation","keywords":"Metadata; Computer science; Data curation; Information retrieval; World Wide Web; Usability; Data science; Database","score_opus":0.0670437112549849,"score_gpt":0.3748008608905625,"score_spread":0.30775714963557765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411622822","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.095107995,0.0026521648,0.7900223,0.0023526321,0.00031572118,0.0012211278,0.04822694,0.054238692,0.005862403],"genre_scores_gemma":[0.121903695,0.00089252746,0.7748445,0.00097376853,0.00011369149,0.0011135027,0.09583174,0.002228538,0.0020979848],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.994651,0.0016261328,0.0009918237,0.0012584662,0.0012541056,0.00021837799],"domain_scores_gemma":[0.98528254,0.008973144,0.001028949,0.002495877,0.0019721533,0.0002473888],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005180649,0.0018701017,0.0012915744,0.0057198647,0.0009539151,0.003143613,0.001981616,0.0011259575,0.002434405],"category_scores_gemma":[0.020395117,0.0005544748,0.0028615082,0.0028855668,0.00069353543,0.004349273,0.0036674999,0.002230822,0.0032056114],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010318462,0.0007877296,0.029731626,0.0064410204,0.0006110443,0.0021907166,0.0037645618,0.022713464,0.15484676,0.020438321,0.07248763,0.6849553],"study_design_scores_gemma":[0.00028994822,0.0004129044,0.021045,0.00085081824,0.00077425264,0.0017111229,0.0030777927,0.51670116,0.15616374,0.04409119,0.25446,0.0004221213],"about_ca_topic_score_codex":0.007058961,"about_ca_topic_score_gemma":0.01327192,"teacher_disagreement_score":0.007058961,"about_ca_system_score_codex":0.0013455207,"about_ca_system_score_gemma":0.004080608,"threshold_uncertainty_score":0.027398169},"labels":[],"label_agreement":null},{"id":"W4411789579","doi":"10.31219/osf.io/94mex_v1","title":"Natural Language Querying of Biological Databases with Large Language Models","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bioinformatics Solutions (Canada)","funders":"AstraZeneca","keywords":"Computer science; Database; Natural (archaeology); Natural language; Natural language processing; Geography; Archaeology","score_opus":0.03145828793273563,"score_gpt":0.32285844174435596,"score_spread":0.2914001538116203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4411789579","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10552313,0.0014242722,0.8467249,0.002920295,0.00007002003,0.00081201043,0.005101821,0.03236231,0.0050612832],"genre_scores_gemma":[0.37652898,0.0007608219,0.6091114,0.0008601766,0.000055567212,0.0005314114,0.009407907,0.0017000445,0.001043633],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9831217,0.008735133,0.00160482,0.0016949525,0.0044926405,0.00035076647],"domain_scores_gemma":[0.93192136,0.054518905,0.0023660664,0.007457247,0.0033184686,0.00041795612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012566735,0.0014028109,0.0014539985,0.0022650969,0.001010223,0.004383065,0.003319756,0.0018060524,0.0025328863],"category_scores_gemma":[0.06612356,0.0009245314,0.002227337,0.0032673297,0.0020260832,0.0096002165,0.004733614,0.0026996543,0.001085528],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0023514132,0.0015788538,0.019157125,0.007271505,0.0014765324,0.001637567,0.0072858916,0.24219903,0.08663272,0.13933955,0.038760822,0.45230895],"study_design_scores_gemma":[0.00019210235,0.00018244554,0.0013732662,0.00014979219,0.00016386551,0.00041278804,0.0012580756,0.84732515,0.039970066,0.08825514,0.020616386,0.00010091255],"about_ca_topic_score_codex":0.0069330554,"about_ca_topic_score_gemma":0.010371597,"teacher_disagreement_score":0.012566735,"about_ca_system_score_codex":0.002358359,"about_ca_system_score_gemma":0.0026787515,"threshold_uncertainty_score":0.06646001},"labels":[],"label_agreement":null},{"id":"W4412030078","doi":"10.2196/72133","title":"In Silico Analysis and Validation of A Disintegrin and Metalloprotease (ADAM) 17 Gene Missense Variants: Structural Bioinformatics Study","year":2025,"lang":"en","type":"article","venue":"JMIR Bioinformatics and Biotechnology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"In silico; Preprint; Missense mutation; Computational biology; Biology; Gene; Genetics; Computer science; Mutation; World Wide Web","score_opus":0.007146013235647111,"score_gpt":0.27427687062464917,"score_spread":0.26713085738900205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412030078","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94475454,0.0010295502,0.034522533,0.0005047069,0.00010547843,0.0002992335,0.011601748,0.0039608274,0.0032212068],"genre_scores_gemma":[0.8908611,0.00067443296,0.08086148,0.00017507221,0.000033425196,0.00027060896,0.025478616,0.00067891215,0.00096636766],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994574,0.00016397343,0.000053877095,0.00014834604,0.00010532228,0.00007113642],"domain_scores_gemma":[0.99927205,0.00042167117,0.000079769474,0.00004316389,0.00010600867,0.00007729835],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016377814,0.0011653461,0.001413147,0.0012957335,0.0007344584,0.0008736523,0.0010940613,0.00072385726,0.0027112516],"category_scores_gemma":[0.0018275899,0.00033683045,0.0018162101,0.0011873998,0.00028191364,0.000367934,0.00052111794,0.0006382354,0.0012860971],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0072055105,0.0034681545,0.12496747,0.00504995,0.0024804259,0.014802354,0.0013129615,0.43195093,0.2622406,0.0077806194,0.028594509,0.110146634],"study_design_scores_gemma":[0.00065697794,0.0015136505,0.026592748,0.00017308131,0.0011945252,0.0027905137,0.00041562377,0.91413575,0.036999825,0.002276783,0.013136069,0.000114446135],"about_ca_topic_score_codex":0.0014515932,"about_ca_topic_score_gemma":0.0028470897,"teacher_disagreement_score":0.0027112516,"about_ca_system_score_codex":0.00047610656,"about_ca_system_score_gemma":0.0012431608,"threshold_uncertainty_score":0.009070098},"labels":[],"label_agreement":null},{"id":"W4412163702","doi":"10.1158/1557-3265.aimachine-pr-05","title":"Abstract PR-05: Learning the Language of Somatic Mutations: A Large Language Model Approach to Precision Oncology","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Precision oncology; Somatic cell; Oncology; Medicine; Computational biology; Language model; Natural language processing; Internal medicine; Computer science; Artificial intelligence; Biology; Genetics; Cancer; Gene","score_opus":0.13813087850948066,"score_gpt":0.5422115612473367,"score_spread":0.404080682737856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163702","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09125699,0.0009942508,0.8897695,0.0044158404,0.00020375881,0.00015138084,0.0023471413,0.007580348,0.003280854],"genre_scores_gemma":[0.7638505,0.00041764337,0.22234128,0.0013700766,0.00024968537,0.00023892132,0.005031324,0.0004783567,0.0060221367],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990847,0.0004117967,0.00003954002,0.00030731852,0.00009985854,0.00005684843],"domain_scores_gemma":[0.9963983,0.0026979418,0.00015454624,0.00028870886,0.00033936367,0.000121057645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020805032,0.000896115,0.0006775644,0.00078031357,0.00041189027,0.0011363968,0.0015976643,0.001081826,0.0033356387],"category_scores_gemma":[0.005732516,0.0004085092,0.0013969224,0.0005872118,0.00054142583,0.0015318057,0.0012832383,0.002794811,0.0009451538],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005082619,0.00037354618,0.0047151116,0.00022236073,0.00028138453,0.00035356352,0.00031892798,0.60998887,0.005733816,0.0177134,0.023051618,0.33673918],"study_design_scores_gemma":[0.000013869078,0.000032577587,0.00016041797,0.00000708528,0.000013473379,0.000020924414,0.000011606386,0.9909616,0.0004966686,0.0076475767,0.0006279564,0.0000061837463],"about_ca_topic_score_codex":0.009736247,"about_ca_topic_score_gemma":0.009766894,"teacher_disagreement_score":0.009736247,"about_ca_system_score_codex":0.0013250558,"about_ca_system_score_gemma":0.0013741251,"threshold_uncertainty_score":0.019359112},"labels":[],"label_agreement":null},{"id":"W4412163719","doi":"10.1158/1557-3265.aimachine-b015","title":"Abstract B015: Large language models to predict cancer risk from free-text clinical notes","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Sunnybrook Health Science Centre; University Health Network; University of Toronto; North Pacific Marine Science Organization; Princess Margaret Cancer Centre","funders":"","keywords":"Cancer; Medicine; Linguistics; Computer science; Internal medicine; Philosophy","score_opus":0.17617279812456935,"score_gpt":0.5355322737544392,"score_spread":0.3593594756298699,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163719","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6067453,0.0046082535,0.17824705,0.010734419,0.00091838907,0.001398695,0.16559908,0.023736207,0.008012711],"genre_scores_gemma":[0.7743015,0.0007858332,0.09220918,0.0011671366,0.0004114502,0.000802958,0.12357484,0.00032258005,0.0064244666],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9990264,0.0004126426,0.000073516334,0.00030602724,0.00010265616,0.00007875028],"domain_scores_gemma":[0.9932541,0.005200262,0.00025808983,0.00042364185,0.0006372391,0.00022659094],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031248678,0.001510074,0.0006210192,0.0015034361,0.00062581804,0.0013394622,0.0020517176,0.0011017454,0.0052840076],"category_scores_gemma":[0.014528538,0.0005244246,0.0014580347,0.00087900343,0.00034765125,0.0011070404,0.0012428724,0.0016634744,0.002457957],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0040309066,0.0013868226,0.19627285,0.0011057487,0.0012265429,0.0010124535,0.000730479,0.28488824,0.0030019726,0.004105739,0.14664492,0.35559332],"study_design_scores_gemma":[0.00026784075,0.00018462486,0.011941304,0.000099188626,0.00013788929,0.00009652994,0.00012634983,0.9761326,0.0007827233,0.0042502754,0.0059254398,0.000055265045],"about_ca_topic_score_codex":0.10686603,"about_ca_topic_score_gemma":0.13098317,"teacher_disagreement_score":0.10686603,"about_ca_system_score_codex":0.0024645433,"about_ca_system_score_gemma":0.0026438965,"threshold_uncertainty_score":0.212488},"labels":[],"label_agreement":null},{"id":"W4412163721","doi":"10.1158/1557-3265.aimachine-a028","title":"Abstract A028: Automated Rule Synthesis from Literature for Agent-Based Modeling of the Tumor Microenvironment","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Tumor microenvironment; Cancer research; Computer science; Medicine; Computational biology; Tumor cells; Biology","score_opus":0.10806087661033964,"score_gpt":0.4620204237970602,"score_spread":0.35395954718672057,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163721","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0135046225,0.0012193756,0.9077352,0.0009719787,0.000250736,0.0008317351,0.010995344,0.0580242,0.00646682],"genre_scores_gemma":[0.08089612,0.0007252636,0.8920637,0.00029000378,0.000091202055,0.00070711423,0.020081725,0.0019574657,0.0031874552],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9978005,0.0007612871,0.0002519486,0.0004387059,0.000667025,0.00008043628],"domain_scores_gemma":[0.9932789,0.004442578,0.0003195491,0.0007878112,0.0009956465,0.00017541635],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0031277684,0.0013439059,0.000695411,0.0039017105,0.0007111084,0.0026106478,0.0023748418,0.0011301398,0.013461897],"category_scores_gemma":[0.013926385,0.0005662396,0.0023297577,0.0012711599,0.0007556709,0.001461167,0.0025580795,0.0010404444,0.0040798895],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00052092626,0.0004367562,0.004335792,0.004377961,0.00048605844,0.0019369798,0.0008511035,0.2660973,0.025661835,0.06597218,0.071678735,0.5576444],"study_design_scores_gemma":[0.00015734074,0.00013796384,0.0004786331,0.00032585688,0.0001373437,0.00046945992,0.00014559552,0.83659214,0.01702954,0.045750942,0.098706394,0.00006872639],"about_ca_topic_score_codex":0.0044650724,"about_ca_topic_score_gemma":0.0074707028,"teacher_disagreement_score":0.013461897,"about_ca_system_score_codex":0.0009090789,"about_ca_system_score_gemma":0.0027054008,"threshold_uncertainty_score":0.045034528},"labels":[],"label_agreement":null},{"id":"W4412163772","doi":"10.1158/1557-3265.aimachine-a018","title":"Abstract A018: Accelerating drug discovery at an HBCU with AI/ML: Text mining, computational modeling, and drug repurposing approaches","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Drug repositioning; Drug; Repurposing; Drug discovery; Computer science; Medicine; Pharmacology; Bioinformatics; Engineering; Biology","score_opus":0.22641333346967343,"score_gpt":0.4743242718112853,"score_spread":0.24791093834161185,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163772","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10073698,0.0038734328,0.8107602,0.02026491,0.00069805887,0.0008186115,0.0069009867,0.031620894,0.024325935],"genre_scores_gemma":[0.23212439,0.0028543219,0.74835795,0.0015252273,0.00039294641,0.0003521652,0.005427094,0.00077727635,0.008188609],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99938226,0.00019539519,0.000041648484,0.00010794144,0.00024138237,0.000031358246],"domain_scores_gemma":[0.99735415,0.0014087645,0.00022135137,0.00024225061,0.00061833556,0.00015508963],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016707472,0.0009107539,0.0005588244,0.0022862614,0.00058840966,0.001965288,0.0009286773,0.00057518075,0.0071236813],"category_scores_gemma":[0.003507298,0.0003447513,0.00085012853,0.001523922,0.000440706,0.001455153,0.00082170137,0.0012645845,0.001972712],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046174423,0.0006100883,0.008709525,0.0009841134,0.00024529095,0.00041368272,0.00018808807,0.14649272,0.041886434,0.02894266,0.08768767,0.68337804],"study_design_scores_gemma":[0.0000957241,0.0002175387,0.0017440146,0.000094768446,0.000072335766,0.00011272313,0.000092404756,0.92108446,0.01777709,0.022726718,0.03593386,0.000048393802],"about_ca_topic_score_codex":0.00599818,"about_ca_topic_score_gemma":0.0076617226,"teacher_disagreement_score":0.0071236813,"about_ca_system_score_codex":0.0009678791,"about_ca_system_score_gemma":0.0023712947,"threshold_uncertainty_score":0.02383107},"labels":[],"label_agreement":null},{"id":"W4412163796","doi":"10.1158/1557-3265.aimachine-a006","title":"Abstract A006: Data Curation and Knowledge Integration Pipeline for Biomarker Discovery","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Ontario Institute for Cancer Research","funders":"","keywords":"Data curation; Pipeline (software); Data science; Biomarker; Medicine; Data integration; Biomarker discovery; Computational biology; Computer science; Biology; Data mining; Proteomics","score_opus":0.41854907872434727,"score_gpt":0.6062314006747581,"score_spread":0.18768232195041085,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163796","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0021185684,0.0005430286,0.36750063,0.001055304,0.00028837982,0.0016858499,0.18459387,0.43719745,0.005016972],"genre_scores_gemma":[0.015207541,0.00067550386,0.5888225,0.0009571126,0.00013504628,0.0026544102,0.36816674,0.018065242,0.0053158193],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9972146,0.0004377783,0.00043504933,0.00088542857,0.0008689419,0.00015810916],"domain_scores_gemma":[0.99440384,0.0019062466,0.00035052447,0.001336809,0.0015519809,0.00045061015],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004807425,0.004133721,0.0021704033,0.008657182,0.0014739599,0.0052445554,0.003765478,0.0015961796,0.0633984],"category_scores_gemma":[0.0136929555,0.0016752931,0.00351138,0.0058163935,0.00061226427,0.0034222945,0.005252362,0.0030774802,0.051563777],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012533324,0.0004657638,0.004435375,0.004140853,0.0007474146,0.0007477604,0.0004954059,0.00667609,0.021493986,0.007289386,0.688295,0.26395968],"study_design_scores_gemma":[0.0012174934,0.00051402627,0.012315242,0.00090380525,0.00058540225,0.00094297255,0.0005181322,0.19274189,0.050262224,0.061232653,0.67826796,0.0004982303],"about_ca_topic_score_codex":0.0074625853,"about_ca_topic_score_gemma":0.007836832,"teacher_disagreement_score":0.0633984,"about_ca_system_score_codex":0.0011809185,"about_ca_system_score_gemma":0.0053472286,"threshold_uncertainty_score":0.21208876},"labels":[],"label_agreement":null},{"id":"W4412163803","doi":"10.1158/1557-3265.aimachine-b006","title":"Abstract B006: Using large language models for scalable extraction of real-world progression events across multiple cancer types","year":2025,"lang":"en","type":"article","venue":"Clinical Cancer Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cancer; Computer science; Scalability; Extraction (chemistry); Medicine; Internal medicine; Chemistry; Database","score_opus":0.2221175994035013,"score_gpt":0.6034192008268803,"score_spread":0.381301601423379,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412163803","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09869745,0.002182539,0.777882,0.0035911566,0.0005814385,0.0011246953,0.04445817,0.06784624,0.0036362733],"genre_scores_gemma":[0.26440662,0.0005414493,0.663772,0.00083402224,0.00019835627,0.0009090314,0.0651534,0.0012637863,0.0029212302],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9971871,0.001168497,0.00034459968,0.0008145488,0.00036494926,0.000120314675],"domain_scores_gemma":[0.98912597,0.007554626,0.00063323363,0.0012860226,0.0011413587,0.00025884554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005256624,0.0017622001,0.0006679441,0.0023008129,0.0007658226,0.0022203636,0.001557189,0.0012056251,0.005310883],"category_scores_gemma":[0.021267084,0.00061708095,0.002672448,0.0014189591,0.00053298345,0.0028231791,0.0024794224,0.0022813547,0.004299686],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0020687582,0.000670922,0.043090638,0.002509259,0.0008607274,0.0011016019,0.0017965285,0.0963306,0.02139459,0.007917749,0.09217531,0.7300833],"study_design_scores_gemma":[0.00027156484,0.0003071756,0.008315811,0.00021877485,0.00022546866,0.00040696288,0.000547138,0.9288609,0.013721396,0.01851787,0.028464004,0.00014286776],"about_ca_topic_score_codex":0.013969025,"about_ca_topic_score_gemma":0.02105846,"teacher_disagreement_score":0.013969025,"about_ca_system_score_codex":0.0011029562,"about_ca_system_score_gemma":0.0025206795,"threshold_uncertainty_score":0.027800024},"labels":[],"label_agreement":null},{"id":"W4412228289","doi":"","title":"Repurposing electronic health data for clinical research and beyond:Experiences from a specialised neurorehabilitation clinic treating patients with acquired brain injury","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Spinal Cord Injury BC","funders":"","keywords":"Neurorehabilitation; Repurposing; Acquired brain injury; Medicine; Physical medicine and rehabilitation; Psychology; Physical therapy; Rehabilitation; Engineering","score_opus":0.13909476938132462,"score_gpt":0.4749362528105395,"score_spread":0.33584148342921494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412228289","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8898703,0.010783768,0.005419778,0.072408676,0.00066778937,0.0010629838,0.001468977,0.00022222528,0.018095491],"genre_scores_gemma":[0.9472583,0.010880666,0.013984148,0.020693088,0.0007635414,0.00057737675,0.0014863381,0.0002638794,0.0040925024],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9422383,0.04454453,0.0034747252,0.0017077436,0.0058491617,0.002185599],"domain_scores_gemma":[0.86253357,0.09248005,0.009068449,0.0057583884,0.011052858,0.019106708],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.046970457,0.00044821072,0.00071756355,0.0013205727,0.003701734,0.008305302,0.002839111,0.002590384,0.0049812463],"category_scores_gemma":[0.1423709,0.00062690116,0.0010097697,0.0024217085,0.004171645,0.006630275,0.010629499,0.0042440086,0.0016624581],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010054403,0.0033514549,0.106145166,0.0036017601,0.00031237814,0.013605514,0.47078744,0.0007570558,0.0014065607,0.0036942286,0.06077635,0.33455667],"study_design_scores_gemma":[0.0004383462,0.0042401333,0.08392074,0.006499386,0.00026628491,0.018460628,0.65395176,0.0023343998,0.00159942,0.00399064,0.22383368,0.00046462286],"about_ca_topic_score_codex":0.003820674,"about_ca_topic_score_gemma":0.007874669,"teacher_disagreement_score":0.046970457,"about_ca_system_score_codex":0.002829369,"about_ca_system_score_gemma":0.008437161,"threshold_uncertainty_score":0.24840647},"labels":[],"label_agreement":null},{"id":"W4412345475","doi":"10.31219/osf.io/6dgv7_v1","title":"Open Sharing of Neuroscience Data in the Canadian Context","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"University of Toronto","keywords":"Context (archaeology); Open science; Data sharing; Neuroscience; Computer science; Data science; Psychology; Cognitive science; Geography; Physics; Medicine","score_opus":0.13616575579052242,"score_gpt":0.38077413683314565,"score_spread":0.24460838104262322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412345475","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10243215,0.01575692,0.09040717,0.39371032,0.003000151,0.004297653,0.04260402,0.001714863,0.34607664],"genre_scores_gemma":[0.6939318,0.020605475,0.16428097,0.039928574,0.0011301687,0.0031801336,0.02255345,0.0010159143,0.053373467],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.92775667,0.027907038,0.00556349,0.0060973926,0.025204815,0.0074706604],"domain_scores_gemma":[0.7639541,0.08206993,0.011604953,0.03626949,0.08398433,0.022117244],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.06979218,0.00062709063,0.00092429837,0.008340755,0.022097776,0.016889274,0.0050297384,0.002619741,0.010074791],"category_scores_gemma":[0.17042442,0.0007004521,0.0010554771,0.020421002,0.011457369,0.009866104,0.018133068,0.003498353,0.0012459239],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007359499,0.00012753242,0.026275491,0.0022084385,0.0003425236,0.0020888287,0.05147657,0.002386951,0.0023630871,0.41383514,0.24425307,0.25390637],"study_design_scores_gemma":[0.00009548898,0.00005391933,0.030539028,0.0021841556,0.00012638091,0.0005231409,0.024144195,0.0017094974,0.001368185,0.09483628,0.8440636,0.00035605836],"about_ca_topic_score_codex":0.9491069,"about_ca_topic_score_gemma":0.9537294,"teacher_disagreement_score":0.99497026,"about_ca_system_score_codex":0.093682535,"about_ca_system_score_gemma":0.39435866,"threshold_uncertainty_score":0.67971754},"labels":[],"label_agreement":null},{"id":"W4412521978","doi":"10.1136/bmjebm-2025-113715","title":"Co-production and implementation of an evidence collation strategy for a novel point-of-care information resource: gpevidence.org","year":2025,"lang":"en","type":"article","venue":"BMJ evidence-based medicine","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"National Institute for Health and Care Research","keywords":"Collation; Production (economics); Resource (disambiguation); Point (geometry); Information resource; Computer science; Business; Knowledge management; Mathematics; Economics","score_opus":0.07713404106508719,"score_gpt":0.4283125794856482,"score_spread":0.351178538420561,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412521978","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.016910229,0.003506443,0.6462377,0.11919771,0.008517009,0.0230829,0.05949875,0.046981808,0.0760675],"genre_scores_gemma":[0.047798418,0.001310749,0.87702405,0.007734573,0.002224097,0.006506118,0.036863085,0.00714143,0.013397494],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9024421,0.04754167,0.021043764,0.0069606407,0.020168077,0.0018438869],"domain_scores_gemma":[0.5006598,0.26858705,0.024115546,0.07906412,0.113717966,0.013855616],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.13448413,0.0020271249,0.002775159,0.025949366,0.0027677643,0.012031636,0.004761803,0.0048266947,0.06316456],"category_scores_gemma":[0.4032358,0.0018657481,0.0033754797,0.017724864,0.0015183946,0.011043658,0.017267222,0.005219928,0.03252321],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009270587,0.00047433545,0.008132983,0.005422674,0.00069533545,0.0020086428,0.0035576248,0.0013893388,0.0038366637,0.008943129,0.44485143,0.51976097],"study_design_scores_gemma":[0.0011929739,0.0004893355,0.009349131,0.008915476,0.0013775195,0.001990768,0.003005616,0.013516929,0.014572475,0.037755158,0.90728176,0.0005528371],"about_ca_topic_score_codex":0.0052128113,"about_ca_topic_score_gemma":0.0066269627,"teacher_disagreement_score":0.8655159,"about_ca_system_score_codex":0.004076769,"about_ca_system_score_gemma":0.034504537,"threshold_uncertainty_score":0.71122855},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"methods","study_design":"not_applicable","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W4412567625","doi":"10.31219/osf.io/6dgv7_v2","title":"Open Sharing of Neuroscience Data in the Canadian Context","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"University of Toronto","keywords":"Context (archaeology); Open science; Data sharing; Open data; Cognitive science; Neuroscience; Data science; Psychology; Computer science; World Wide Web; Geography; Physics; Medicine","score_opus":0.13616575579052242,"score_gpt":0.38077413683314565,"score_spread":0.24460838104262322,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412567625","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10522894,0.016517151,0.087069005,0.40882742,0.003135168,0.0045940364,0.04135333,0.0016517433,0.33162314],"genre_scores_gemma":[0.69597423,0.020382194,0.16026612,0.04421565,0.0011178486,0.0034317605,0.021113308,0.0009555951,0.05254334],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.92349863,0.029718727,0.005828488,0.006435507,0.026431689,0.008087002],"domain_scores_gemma":[0.74888873,0.083397195,0.012404841,0.038640745,0.09203367,0.024634756],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.07401607,0.0006327339,0.00095877255,0.008211122,0.022852745,0.016776513,0.0053775725,0.0026594186,0.009855899],"category_scores_gemma":[0.17433785,0.00073416386,0.001079754,0.019265205,0.011739321,0.009515696,0.018661845,0.0036660973,0.0011916516],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008018424,0.00013407232,0.028671082,0.0023069466,0.00038448596,0.0021315766,0.055783145,0.0023862044,0.0025333578,0.40294155,0.24865286,0.25327286],"study_design_scores_gemma":[0.000101358,0.00006128508,0.034033246,0.002330508,0.00013853043,0.00053162314,0.02493321,0.0016452012,0.0013533268,0.08849126,0.84599805,0.00038235492],"about_ca_topic_score_codex":0.9549677,"about_ca_topic_score_gemma":0.9614479,"teacher_disagreement_score":0.9946224,"about_ca_system_score_codex":0.09879288,"about_ca_system_score_gemma":0.41528234,"threshold_uncertainty_score":0.71679586},"labels":[],"label_agreement":null},{"id":"W4412964009","doi":"10.1109/iccc64910.2025.11077215","title":"Improving Discharge Summary Generation through Clinical Text Summarization with QLoRA and LLMs","year":2025,"lang":"","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing","score_opus":0.02415710483460701,"score_gpt":0.313254582083443,"score_spread":0.28909747724883594,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412964009","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04381196,0.0028019997,0.73199177,0.0033628677,0.0009905315,0.001290531,0.070797004,0.13942173,0.005531608],"genre_scores_gemma":[0.18974927,0.00095271145,0.69541687,0.0009951419,0.000534929,0.0007579715,0.10394242,0.0024200857,0.0052305916],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981578,0.00048775857,0.00037304295,0.00044457006,0.00043054108,0.000106274],"domain_scores_gemma":[0.99338526,0.0034917318,0.0005479946,0.0006445763,0.0017304186,0.00020002025],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022015255,0.0015246621,0.0011356967,0.0056490344,0.0006822987,0.0023958154,0.0010834665,0.0009777755,0.012178778],"category_scores_gemma":[0.013249366,0.00046444434,0.0015013717,0.00308378,0.00023020821,0.0019446226,0.0020477527,0.0011321882,0.0064001386],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008580993,0.00029238436,0.007561464,0.001959266,0.0004320814,0.00070160365,0.00058316166,0.011986414,0.021482952,0.0028694058,0.12218452,0.8290887],"study_design_scores_gemma":[0.0006526063,0.00072086905,0.015913555,0.0007316944,0.0013650219,0.0012433118,0.0015336975,0.6870105,0.079616964,0.029777583,0.18116276,0.0002714248],"about_ca_topic_score_codex":0.0038932557,"about_ca_topic_score_gemma":0.0055944407,"teacher_disagreement_score":0.012178778,"about_ca_system_score_codex":0.0006050279,"about_ca_system_score_gemma":0.0023059763,"threshold_uncertainty_score":0.0407421},"labels":[],"label_agreement":null},{"id":"W4412979870","doi":"10.1016/j.jclinepi.2025.111921","title":"GRADE concept paper 9: rationale and process for creating a GRADE Ontology","year":2025,"lang":"en","type":"article","venue":"Journal of Clinical Epidemiology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hamilton Health Sciences; Bruyère; Western University; University of Alberta; McMaster University; Impact","funders":"","keywords":"Ontology; Computer science; Process (computing); Medicine; Information retrieval; Epistemology; Philosophy; Programming language","score_opus":0.19519089316649038,"score_gpt":0.5127765105814261,"score_spread":0.31758561741493574,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4412979870","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005724424,0.0013308047,0.67817605,0.093290895,0.008804131,0.041387007,0.0063005006,0.002136434,0.16284974],"genre_scores_gemma":[0.02842259,0.000902228,0.89299047,0.012269488,0.0007624033,0.022360401,0.0032136524,0.0010577126,0.038020983],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8903685,0.048599754,0.019397765,0.0074224407,0.02994491,0.0042666337],"domain_scores_gemma":[0.73606193,0.098288305,0.012289523,0.03295535,0.109345384,0.01105947],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1512818,0.00144769,0.0017572532,0.009345011,0.008809652,0.027107656,0.010254747,0.017905284,0.03374501],"category_scores_gemma":[0.29619467,0.0028769758,0.0052896393,0.0068059135,0.015426486,0.019738102,0.017149586,0.01918718,0.021144379],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013656933,0.00015423514,0.00073816016,0.0018194411,0.000040834002,0.000511916,0.011921831,0.0009596362,0.0010116508,0.8275108,0.09816023,0.057034682],"study_design_scores_gemma":[0.00015065576,0.00008887792,0.0005764198,0.0031326197,0.000089362366,0.0002736951,0.0040046205,0.0013443729,0.0018372402,0.13497418,0.85339576,0.00013223739],"about_ca_topic_score_codex":0.018664028,"about_ca_topic_score_gemma":0.021606617,"teacher_disagreement_score":0.84871817,"about_ca_system_score_codex":0.019089086,"about_ca_system_score_gemma":0.0761745,"threshold_uncertainty_score":0.8000642},"labels":[],"label_agreement":null},{"id":"W4413081669","doi":"10.7554/elife.105565.3.sa0","title":"eLife Assessment: Interpretable protein-DNA interactions captured by structure-sequence optimization","year":2025,"lang":"en","type":"peer-review","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Queen's University","funders":"","keywords":"Sequence (biology); Computational biology; DNA; Computer science; Artificial intelligence; Biology; Genetics","score_opus":0.017654183877419308,"score_gpt":0.3434417609449189,"score_spread":0.3257875770674996,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413081669","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09187037,0.0031234995,0.7165667,0.004433936,0.0012960843,0.0010746649,0.04987471,0.10232073,0.029439297],"genre_scores_gemma":[0.2702272,0.0016852277,0.62312275,0.0007973706,0.00028057318,0.0006907842,0.07453047,0.007497236,0.02116835],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9956637,0.0009996009,0.00027790444,0.0006254311,0.0022491494,0.00018434168],"domain_scores_gemma":[0.99077976,0.0042902557,0.00061272143,0.0010237888,0.0030842877,0.00020915159],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005854471,0.0012893154,0.001019275,0.003925418,0.0006751347,0.0026850954,0.002200021,0.0013241507,0.013678319],"category_scores_gemma":[0.02565851,0.00038364256,0.0009542656,0.0025594996,0.0007787072,0.00208209,0.0017440958,0.001355148,0.0070877452],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011333382,0.00038309762,0.014138553,0.002332583,0.00044025353,0.0011145369,0.00065358425,0.036900826,0.03691656,0.020156581,0.28694445,0.59888566],"study_design_scores_gemma":[0.0002774321,0.00023560303,0.010510594,0.00054585485,0.00039611323,0.00092238537,0.00075267965,0.51096827,0.11095928,0.058852084,0.305428,0.00015181373],"about_ca_topic_score_codex":0.0034283136,"about_ca_topic_score_gemma":0.007284331,"teacher_disagreement_score":0.9941455,"about_ca_system_score_codex":0.0008105951,"about_ca_system_score_gemma":0.0038398472,"threshold_uncertainty_score":0.045758486},"labels":[],"label_agreement":null},{"id":"W4413109704","doi":"10.1093/bioadv/vbaf131","title":"Next generation biobanking ontology: introducing–omics contextual data to biobanking ontology","year":2024,"lang":"en","type":"article","venue":"Bioinformatics Advances","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Genome British Columbia; Simon Fraser University","funders":"King Fahad Medical City","keywords":"Biobank; Ontology; Computer science; Data science; Ontology-based data integration; Data discovery; Open Biomedical Ontologies; Data integration; Information retrieval; Metadata; Data management; World Wide Web; Data mining; Semantic Web; Suggested Upper Merged Ontology; Bioinformatics","score_opus":0.08394483396799426,"score_gpt":0.3364466195746183,"score_spread":0.25250178560662406,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413109704","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0063441736,0.0010180352,0.9446732,0.007913041,0.0006529162,0.00080386084,0.010743006,0.0063170027,0.02153479],"genre_scores_gemma":[0.036835685,0.0017663043,0.92607075,0.002574081,0.00026892897,0.0006850613,0.025184268,0.0011998715,0.005415011],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9952266,0.0012789733,0.00078651635,0.0007370993,0.0016061268,0.00036467318],"domain_scores_gemma":[0.9943731,0.0015923628,0.00051289477,0.0013555306,0.0017095325,0.00045655452],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008143956,0.0006534371,0.00062756235,0.0044865655,0.0017049244,0.004431857,0.0021378633,0.0013172934,0.0026283793],"category_scores_gemma":[0.010325292,0.0005515447,0.0017422048,0.0050582895,0.001868109,0.007986521,0.0054444703,0.0032721628,0.0016882978],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018186777,0.00027543402,0.006235251,0.0018113361,0.00019459345,0.0013394656,0.0040719206,0.0065516564,0.012912291,0.6554057,0.09909571,0.21192479],"study_design_scores_gemma":[0.00002572212,0.000019569308,0.0020504333,0.0008171987,0.000085253305,0.0006755379,0.000740134,0.017247118,0.0047686147,0.12345349,0.85004085,0.00007597222],"about_ca_topic_score_codex":0.02005384,"about_ca_topic_score_gemma":0.019025503,"teacher_disagreement_score":0.02005384,"about_ca_system_score_codex":0.0031887083,"about_ca_system_score_gemma":0.009362637,"threshold_uncertainty_score":0.0430699},"labels":[],"label_agreement":null},{"id":"W4413127818","doi":"10.18280/ts.420443","title":"U-NeuroSegNet: A Deep Learning NIWatershed Based Data-Driven Panoptic Segmentation Framework for Identifying Specific Conditions in Neurodegenerative Neurological Disorders","year":2025,"lang":"en","type":"article","venue":"Traitement du signal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Panopticon; Segmentation; Artificial intelligence; Computer science; Deep learning; Neuroscience; Psychology; Sociology","score_opus":0.04142225170417025,"score_gpt":0.3220101944899661,"score_spread":0.28058794278579585,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413127818","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05439647,0.0016776124,0.92720777,0.0008612202,0.00025981275,0.00025463547,0.0027535087,0.008356787,0.0042322436],"genre_scores_gemma":[0.56553346,0.0013556449,0.40855366,0.0011422414,0.00016666044,0.0005598236,0.008485299,0.00053333835,0.013669857],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9998386,0.000022382805,0.000011585676,0.000058673744,0.000036997513,0.0000318058],"domain_scores_gemma":[0.99987936,0.000030131328,0.000016303537,0.000015433521,0.00004480453,0.000013894046],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00039496346,0.0012416608,0.00072414055,0.0010399162,0.00038917921,0.0007665283,0.0013090987,0.0011962679,0.002609075],"category_scores_gemma":[0.00087550184,0.00043110925,0.0009391682,0.0006117445,0.0004355332,0.0008153858,0.0011099022,0.0009309314,0.00074422004],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055426324,0.000245822,0.008212491,0.0002709308,0.0003014682,0.0005806276,0.00015134792,0.3965049,0.017384414,0.009861932,0.026733657,0.53919816],"study_design_scores_gemma":[0.000013523802,0.0000467148,0.0007545778,0.000025291874,0.000022133932,0.00009867092,0.000023852326,0.98579144,0.0050254907,0.0048269536,0.0033577017,0.000013774973],"about_ca_topic_score_codex":0.020437626,"about_ca_topic_score_gemma":0.034010068,"teacher_disagreement_score":0.020437626,"about_ca_system_score_codex":0.0011487522,"about_ca_system_score_gemma":0.0016660573,"threshold_uncertainty_score":0.040637314},"labels":[],"label_agreement":null},{"id":"W4413209346","doi":"10.2196/76870","title":"Building a Standardized Cancer Synoptic Report With Semantic and Syntactic Interoperability: Development Study Using SNOMED CT and Fast Healthcare Interoperability Resources (FHIR)","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Ontario Tobacco Research Unit","funders":"","keywords":"SNOMED CT; Interoperability; Computer science; Semantic interoperability; Preprint; Natural language processing; Information retrieval; Semantics (computer science); Medicine; Artificial intelligence; World Wide Web; Terminology; Programming language; Linguistics","score_opus":0.019342883914030082,"score_gpt":0.34347994256071285,"score_spread":0.32413705864668274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413209346","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46141878,0.0010666106,0.48474443,0.002000341,0.00028347396,0.016731067,0.011882097,0.0077354633,0.014137747],"genre_scores_gemma":[0.16006443,0.0004766521,0.81864,0.0002047899,0.000040981537,0.00316666,0.015351969,0.00073615956,0.0013183003],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9797727,0.010551439,0.0028581424,0.0021583878,0.0043528792,0.00030644255],"domain_scores_gemma":[0.9035076,0.04887477,0.005452847,0.012900421,0.027677065,0.0015873382],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.052155234,0.0007962764,0.000473911,0.0059885955,0.0010753645,0.0036530974,0.0027089447,0.0009549459,0.0025277631],"category_scores_gemma":[0.08932884,0.00078452955,0.0019115574,0.003948425,0.001610559,0.0062946547,0.0042311493,0.0011828031,0.0010636958],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010186007,0.002563007,0.065479144,0.0059779705,0.00036903954,0.0020367913,0.03729351,0.015728176,0.04010853,0.024463661,0.0243887,0.78057283],"study_design_scores_gemma":[0.0010622559,0.006975167,0.1487255,0.007587861,0.0022727763,0.006442316,0.056062676,0.15824437,0.19532976,0.018617189,0.3974401,0.0012400858],"about_ca_topic_score_codex":0.008174414,"about_ca_topic_score_gemma":0.0068272226,"teacher_disagreement_score":0.052155234,"about_ca_system_score_codex":0.002639945,"about_ca_system_score_gemma":0.009024926,"threshold_uncertainty_score":0.27582657},"labels":[],"label_agreement":null},{"id":"W4413290922","doi":"10.1038/s42256-025-01088-6","title":"Boosting the predictive power of protein representations with a corpus of text annotations","year":2025,"lang":"en","type":"article","venue":"Nature Machine Intelligence","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; University of Toronto","funders":"Vector Institute; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Boosting (machine learning); Predictive power; Computer science; Natural language processing; Artificial intelligence; Physics","score_opus":0.005744969379399131,"score_gpt":0.29586281015831767,"score_spread":0.29011784077891856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413290922","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7157468,0.0083516985,0.24438271,0.004515803,0.0009242984,0.00029121936,0.008004234,0.007055377,0.010727787],"genre_scores_gemma":[0.9134581,0.0022490197,0.06791436,0.0006650014,0.000693068,0.00017148584,0.011976135,0.00028767542,0.002585141],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988226,0.00035047642,0.000074638025,0.00031779494,0.00033590652,0.00009853394],"domain_scores_gemma":[0.98517716,0.011474197,0.0005609872,0.0011814148,0.001363028,0.00024315929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033049833,0.0014964802,0.0011344078,0.0045644473,0.00065388717,0.0019790726,0.0012326292,0.001929639,0.0014083532],"category_scores_gemma":[0.01795885,0.00046902924,0.0008991041,0.0036601003,0.0007725415,0.004217129,0.0017464119,0.0024735385,0.0010966638],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028085653,0.0019065187,0.029012362,0.0010310027,0.00076148205,0.00084583985,0.0005198515,0.16664562,0.032774013,0.0070532556,0.034612972,0.72202855],"study_design_scores_gemma":[0.00007632383,0.00018348779,0.003252394,0.000086867214,0.00022584865,0.00014671368,0.00009354602,0.9710547,0.007517559,0.013795236,0.003535646,0.000031718166],"about_ca_topic_score_codex":0.004018251,"about_ca_topic_score_gemma":0.0046655964,"teacher_disagreement_score":0.0045644473,"about_ca_system_score_codex":0.0006163025,"about_ca_system_score_gemma":0.000990389,"threshold_uncertainty_score":0.017478645},"labels":[],"label_agreement":null},{"id":"W4413341034","doi":"10.1038/s41597-026-07298-w","title":"VO: The Vaccine Ontology","year":2025,"lang":"en","type":"article","venue":"Scientific Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Ontario Institute for Cancer Research","funders":"","keywords":"Ontology; Vaccine adjuvant; Compendium; Computer science; Adjuvant; Medicine; DNA vaccination; Immune system; Immunology; Immunization","score_opus":0.03379012968567952,"score_gpt":0.33427289494916207,"score_spread":0.30048276526348255,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413341034","genre_codex":"methods","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008454188,0.0059335986,0.7141721,0.01594586,0.002800804,0.0018268902,0.079233535,0.014468178,0.15716481],"genre_scores_gemma":[0.09016801,0.013307503,0.6908387,0.008611961,0.0012398652,0.002370149,0.15548833,0.003933816,0.034041762],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.996599,0.00075869064,0.00064035464,0.0005483966,0.0011110343,0.0003425415],"domain_scores_gemma":[0.99665415,0.0010458285,0.00036658155,0.00067750487,0.0009146077,0.00034133057],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030923693,0.00093869097,0.00081255933,0.005727623,0.0023114067,0.0049498947,0.0020326637,0.0024655752,0.009039899],"category_scores_gemma":[0.008029468,0.0007261141,0.002014787,0.0058193,0.0020776559,0.0107481545,0.0038882573,0.003169674,0.005040557],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008167126,0.00008875974,0.0014868295,0.0010568971,0.00006806224,0.0003575974,0.0011645783,0.0029304954,0.0025913524,0.7385876,0.13186915,0.11971698],"study_design_scores_gemma":[0.000009752066,0.000012388014,0.00039684065,0.00032113562,0.000022347049,0.00029801208,0.00027744213,0.0021054735,0.00052792707,0.074356616,0.92164946,0.000022629192],"about_ca_topic_score_codex":0.018887002,"about_ca_topic_score_gemma":0.014053475,"teacher_disagreement_score":0.018887002,"about_ca_system_score_codex":0.0030398776,"about_ca_system_score_gemma":0.010175506,"threshold_uncertainty_score":0.037554145},"labels":[],"label_agreement":null},{"id":"W4413467890","doi":"10.2196/68558","title":"Generative Models and Sentence Transformers for the Recognition and Normalization of Continuous and Discontinuous Phenotype Mentions: Model Development and Evaluation","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Preprint; Transformer; Generative grammar; Computer science; Sentence; Artificial intelligence; Coronavirus disease 2019 (COVID-19); Natural language processing; Medicine; Engineering; Pathology; World Wide Web; Electrical engineering","score_opus":0.04340515198957854,"score_gpt":0.3216704328808726,"score_spread":0.27826528089129404,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413467890","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1173194,0.0024694325,0.85228956,0.0017051267,0.00035399612,0.0008548821,0.003564277,0.01647004,0.0049733324],"genre_scores_gemma":[0.6489792,0.0011427272,0.33031508,0.0006856831,0.00015993495,0.0014331636,0.009910889,0.00072277285,0.006650561],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9988356,0.0004810444,0.00007524162,0.00035384716,0.00015201863,0.00010219099],"domain_scores_gemma":[0.99266785,0.006003017,0.00018642542,0.00034226105,0.00068830576,0.00011220413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040700296,0.001988212,0.0012254034,0.001461519,0.0005446694,0.001411078,0.0025056216,0.0019247064,0.0049754],"category_scores_gemma":[0.0104777245,0.00092144107,0.0021773984,0.0008197188,0.0006911222,0.001962939,0.0013344638,0.0034226421,0.0027010518],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00060781755,0.00031991318,0.003773273,0.00029412267,0.0002846098,0.00030008226,0.00023175382,0.7650891,0.0037667332,0.0053206477,0.0068354732,0.21317647],"study_design_scores_gemma":[0.00001303334,0.000027640277,0.00017877802,0.000009033555,0.00001784444,0.00002386721,0.000009047245,0.99733174,0.00069028116,0.0013981014,0.00029390462,0.0000066799457],"about_ca_topic_score_codex":0.023590434,"about_ca_topic_score_gemma":0.02707115,"teacher_disagreement_score":0.023590434,"about_ca_system_score_codex":0.0022142194,"about_ca_system_score_gemma":0.0018998755,"threshold_uncertainty_score":0.046906292},"labels":[],"label_agreement":null},{"id":"W4413491517","doi":"10.1007/978-3-031-97788-6_5","title":"Where Does SNOMED CT® Come From?","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Computer science; Linguistics; Philosophy; Terminology","score_opus":0.00804584144450476,"score_gpt":0.2851437619590805,"score_spread":0.2770979205145757,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491517","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0042689005,0.034913097,0.25323227,0.07817131,0.022095667,0.00019845123,0.016825533,0.013112768,0.577182],"genre_scores_gemma":[0.03641789,0.054845624,0.26676142,0.058382902,0.00829055,0.00028501757,0.03910552,0.019684376,0.51622677],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990397,0.00016979994,0.00009489688,0.00014585999,0.0004968852,0.000052805837],"domain_scores_gemma":[0.9978198,0.00095687504,0.000114435934,0.00021228917,0.00068196724,0.00021471537],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019002534,0.0007271306,0.00057754706,0.0037719733,0.0007538674,0.005976015,0.00096992886,0.0018150797,0.04414335],"category_scores_gemma":[0.006143875,0.00049811765,0.0005970294,0.0052342936,0.0013931235,0.012129844,0.001838717,0.0027029584,0.032650657],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000045314784,0.000015425903,0.00020143468,0.00049637235,0.000012984439,0.00016906638,0.00041250655,0.00024679457,0.0018929932,0.13369723,0.49103913,0.3717707],"study_design_scores_gemma":[0.0000031938494,0.0000055579517,0.00013855963,0.00036088796,0.0000071581044,0.0003073634,0.000118884454,0.00031007288,0.0007669455,0.023338448,0.9746259,0.000016984688],"about_ca_topic_score_codex":0.0074848235,"about_ca_topic_score_gemma":0.010663452,"teacher_disagreement_score":0.04414335,"about_ca_system_score_codex":0.0016237831,"about_ca_system_score_gemma":0.0025393893,"threshold_uncertainty_score":0.1476742},"labels":[],"label_agreement":null},{"id":"W4413491554","doi":"10.1007/978-3-031-97788-6_9","title":"Logical Model","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.016576306052353253,"score_gpt":0.31085378262745866,"score_spread":0.2942774765751054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491554","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0040584058,0.00072437647,0.76542515,0.0066738892,0.00064280664,0.00047843807,0.008340445,0.0041602184,0.20949632],"genre_scores_gemma":[0.1328691,0.0021092505,0.59307635,0.004660297,0.0007396755,0.00112867,0.02372611,0.0015726621,0.24011794],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976617,0.0006102507,0.00019667568,0.00062386354,0.00073123653,0.00017632164],"domain_scores_gemma":[0.99836,0.00063039217,0.00007993728,0.00037226922,0.00046090092,0.00009648808],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001969116,0.00112884,0.0005346383,0.002500991,0.0014737325,0.0058990074,0.0023088632,0.0014368836,0.06329846],"category_scores_gemma":[0.0052675894,0.00066965417,0.0017110108,0.0019756735,0.0016832006,0.010949146,0.0030370625,0.0027697682,0.019454401],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000036545724,0.000055690933,0.00021858326,0.000113620874,0.000022284117,0.00011105672,0.0001417495,0.0014977952,0.00047623375,0.925134,0.026870944,0.045321524],"study_design_scores_gemma":[0.0000302696,0.000020292922,0.00009857696,0.00007895692,0.000051127838,0.0002031613,0.00015598028,0.011882104,0.001362088,0.70380694,0.2822903,0.000020232439],"about_ca_topic_score_codex":0.0048199682,"about_ca_topic_score_gemma":0.0053774705,"teacher_disagreement_score":0.06329846,"about_ca_system_score_codex":0.0022163866,"about_ca_system_score_gemma":0.003043801,"threshold_uncertainty_score":0.21175444},"labels":[],"label_agreement":null},{"id":"W4413491577","doi":"10.1007/978-3-031-97788-6_29","title":"User Interfaces","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Human–computer interaction","score_opus":0.009717824009553502,"score_gpt":0.30219116209530344,"score_spread":0.29247333808574993,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491577","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011510397,0.0014188525,0.07103946,0.0015826556,0.0016043365,0.00096067245,0.06353383,0.16364755,0.69506156],"genre_scores_gemma":[0.007835685,0.0013479426,0.02219058,0.0032012446,0.00047295768,0.0015587038,0.05829491,0.038484097,0.86661386],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99928635,0.00011846226,0.000058059828,0.00013556078,0.0003123596,0.00008933196],"domain_scores_gemma":[0.99859565,0.00041628844,0.000037213642,0.00040255988,0.0003865891,0.00016157072],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011477787,0.0026071365,0.0013367553,0.0028769444,0.00086391484,0.0047660824,0.0023298287,0.002271252,0.77963126],"category_scores_gemma":[0.0053775725,0.00077341986,0.0011736299,0.0027771033,0.00040029862,0.0040732473,0.0046203425,0.0013434336,0.732019],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000070739145,0.000027423379,0.00009216556,0.0002884125,0.000010507989,0.000054957574,0.00009692751,0.00007937276,0.000802572,0.0032169756,0.92256683,0.07269315],"study_design_scores_gemma":[0.000030282963,0.000009949169,0.00018232266,0.00009629579,0.000009901425,0.00007512158,0.000034506378,0.0004881279,0.0007055395,0.0027689096,0.99558014,0.00001892419],"about_ca_topic_score_codex":0.0029978002,"about_ca_topic_score_gemma":0.0039605075,"teacher_disagreement_score":0.77963126,"about_ca_system_score_codex":0.00078728056,"about_ca_system_score_gemma":0.000695896,"threshold_uncertainty_score":0.31432927},"labels":[],"label_agreement":null},{"id":"W4413491578","doi":"10.1007/978-3-031-97788-6_15","title":"Templates","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.010334660603437488,"score_gpt":0.30045449143847874,"score_spread":0.2901198308350413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491578","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025306563,0.0007328869,0.26291192,0.0017550854,0.001758624,0.0009861488,0.09892815,0.10181988,0.5285767],"genre_scores_gemma":[0.031039838,0.001597733,0.15211092,0.002071204,0.00039355273,0.0013001672,0.17825688,0.036785204,0.59644455],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988072,0.00018477667,0.000102153586,0.00034195083,0.00044681173,0.00011701896],"domain_scores_gemma":[0.9986883,0.00032329178,0.000048071004,0.00054665277,0.0003243238,0.0000693929],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008997961,0.0017181237,0.0008082053,0.0016830491,0.00083085336,0.0036831726,0.002420916,0.0013371106,0.2869106],"category_scores_gemma":[0.004715526,0.0009074991,0.0012391608,0.0016377502,0.0005215292,0.0041157124,0.002906029,0.0016034715,0.23562998],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016350181,0.00007261059,0.0005310063,0.00041089026,0.000025604711,0.00018698492,0.00028248964,0.00080148294,0.0020825556,0.08356037,0.7260535,0.18582904],"study_design_scores_gemma":[0.000018175879,0.000009281604,0.00013067202,0.0000603202,0.000010965959,0.00014506388,0.000054524025,0.0011070654,0.0025872914,0.016980013,0.97887987,0.000016790536],"about_ca_topic_score_codex":0.0054400493,"about_ca_topic_score_gemma":0.0059052715,"teacher_disagreement_score":0.2869106,"about_ca_system_score_codex":0.001118436,"about_ca_system_score_gemma":0.0016523178,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4413491718","doi":"10.1007/978-3-031-97788-6_20","title":"Subsets","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.009915531185224568,"score_gpt":0.29917972059534864,"score_spread":0.28926418941012405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491718","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012020492,0.001028457,0.13556018,0.001526393,0.0018346244,0.00081967976,0.050300986,0.017472992,0.7794363],"genre_scores_gemma":[0.08605424,0.0013537632,0.087363586,0.001465933,0.00082303013,0.0011585712,0.13992408,0.010437672,0.67141914],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988343,0.00016505443,0.000072828996,0.0004358476,0.00033148256,0.00016046457],"domain_scores_gemma":[0.9988985,0.00015883902,0.000037020727,0.00047274542,0.00036699907,0.00006582561],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00061476335,0.0016145248,0.00072695565,0.0021734475,0.0014810397,0.003192365,0.0014832189,0.0007448005,0.2452885],"category_scores_gemma":[0.0023945314,0.00049667445,0.0009268474,0.0017152474,0.0006073062,0.003576341,0.0028387913,0.0013121938,0.1574453],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004582695,0.00010907706,0.0011860981,0.0004574944,0.000058573074,0.0002002901,0.00032421417,0.001097684,0.005038583,0.1465586,0.4954029,0.34910828],"study_design_scores_gemma":[0.000032039716,0.00004061455,0.0005071134,0.00009127717,0.000028779985,0.0001978211,0.00019455887,0.0014319881,0.004871395,0.06680286,0.92578363,0.000017913717],"about_ca_topic_score_codex":0.0018177615,"about_ca_topic_score_gemma":0.0020531376,"teacher_disagreement_score":0.2452885,"about_ca_system_score_codex":0.0006554237,"about_ca_system_score_gemma":0.0009174319,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4413491721","doi":"10.1007/978-3-031-97788-6_1","title":"What Is SNOMED CT®?","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"SNOMED CT; Medicine; Computer science; Medical physics; Linguistics; Philosophy; Terminology","score_opus":0.011196496056873628,"score_gpt":0.30588948125119797,"score_spread":0.29469298519432435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491721","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.006198416,0.08010508,0.3088017,0.06382329,0.017149739,0.0002764363,0.01708077,0.014650559,0.49191403],"genre_scores_gemma":[0.040538296,0.10127603,0.47496587,0.051623717,0.00822173,0.00038197727,0.03507614,0.012009299,0.2759069],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.998963,0.00021663164,0.00011970178,0.00014251193,0.0005067197,0.00005144505],"domain_scores_gemma":[0.9972498,0.0013825786,0.00016133419,0.00027646404,0.0006429324,0.0002869846],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018659134,0.0008628194,0.00059636147,0.00582785,0.0005425096,0.004832718,0.0009797781,0.0014919229,0.022828171],"category_scores_gemma":[0.0067403656,0.00040782164,0.00055811467,0.0062129223,0.0014220991,0.008552349,0.0013037329,0.0014782837,0.0145704895],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000042125965,0.00001868277,0.0002576648,0.0006081582,0.000019990728,0.00014615011,0.00026729287,0.00037040003,0.0015316154,0.061979596,0.28514585,0.64961255],"study_design_scores_gemma":[0.0000042072215,0.000009413305,0.00024110278,0.0006371111,0.000015212044,0.0006975858,0.00010632431,0.0007122361,0.0009250602,0.037143044,0.9594854,0.000023347977],"about_ca_topic_score_codex":0.0051189912,"about_ca_topic_score_gemma":0.007113601,"teacher_disagreement_score":0.022828171,"about_ca_system_score_codex":0.0013080301,"about_ca_system_score_gemma":0.002265572,"threshold_uncertainty_score":0.076367795},"labels":[],"label_agreement":null},{"id":"W4413491728","doi":"10.1007/978-3-031-97788-6_27","title":"Terminology Services","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Terminology; Computer science; Linguistics; Philosophy","score_opus":0.008646469326045448,"score_gpt":0.29433424407854547,"score_spread":0.2856877747525,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491728","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002642499,0.0023156554,0.15884991,0.005585665,0.0026284293,0.00088308175,0.110324405,0.056050546,0.66071975],"genre_scores_gemma":[0.025043935,0.004568588,0.14624354,0.0050462037,0.0010590266,0.0009638796,0.29574353,0.02285515,0.4984761],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99789304,0.00037497235,0.0002855862,0.0003298857,0.00091826235,0.0001982137],"domain_scores_gemma":[0.9973363,0.0006335604,0.00011758296,0.0007288909,0.000943847,0.00023972036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018565542,0.0014762536,0.0009779871,0.0073576714,0.001655944,0.007040262,0.002507992,0.0018507685,0.24649942],"category_scores_gemma":[0.009443048,0.00065124495,0.0013646744,0.007823392,0.00076664955,0.0078030643,0.005611497,0.0020434395,0.2444436],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009928015,0.000045328317,0.00033425764,0.0005031967,0.00001789431,0.00017864295,0.00046865086,0.0002453991,0.002866993,0.12416268,0.68459594,0.1864817],"study_design_scores_gemma":[0.000008708472,0.000003672533,0.00009317286,0.00006031254,0.0000068437334,0.0001373665,0.00010816645,0.0003829874,0.00081757875,0.0143302195,0.9840389,0.0000121308485],"about_ca_topic_score_codex":0.0065599065,"about_ca_topic_score_gemma":0.007189569,"teacher_disagreement_score":0.24649942,"about_ca_system_score_codex":0.0018890072,"about_ca_system_score_gemma":0.0031577775,"threshold_uncertainty_score":0.8246227},"labels":[],"label_agreement":null},{"id":"W4413491730","doi":"10.1007/978-3-031-97788-6_12","title":"Concept Model","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.011944621934142571,"score_gpt":0.30397836843038745,"score_spread":0.29203374649624486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491730","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007874974,0.002096553,0.6909702,0.0070407027,0.0012657773,0.0013149255,0.023798147,0.0031928218,0.262446],"genre_scores_gemma":[0.16810666,0.0036685388,0.58746076,0.003031808,0.0006241076,0.0023866787,0.04604574,0.0007114829,0.18796425],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979863,0.00040470125,0.00014687196,0.0006208953,0.00069153105,0.00014959375],"domain_scores_gemma":[0.9988788,0.00041510636,0.00005243719,0.00018790572,0.0003988452,0.00006694338],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012935831,0.001129493,0.00048283086,0.0034471948,0.001054994,0.004240109,0.0019754188,0.0012791606,0.047496773],"category_scores_gemma":[0.0043443018,0.00038722574,0.0014443074,0.0028716568,0.0011521676,0.0069628987,0.0018346276,0.0018567417,0.014411416],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006573129,0.000074843396,0.00051045016,0.00026628308,0.00004135508,0.00030011366,0.00033831233,0.0029850435,0.0007423126,0.8451174,0.044586748,0.10497152],"study_design_scores_gemma":[0.00004698643,0.000032476728,0.00026967155,0.00018221347,0.00007133377,0.00055953825,0.00029926482,0.019744579,0.0015020015,0.47039166,0.50686955,0.000030761046],"about_ca_topic_score_codex":0.009234875,"about_ca_topic_score_gemma":0.0067606703,"teacher_disagreement_score":0.047496773,"about_ca_system_score_codex":0.0020065086,"about_ca_system_score_gemma":0.002827419,"threshold_uncertainty_score":0.15889251},"labels":[],"label_agreement":null},{"id":"W4413491737","doi":"10.1007/978-3-031-97788-6_32","title":"Challenges","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Business","score_opus":0.020591814271305898,"score_gpt":0.3146401800249321,"score_spread":0.2940483657536262,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491737","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018643058,0.0037105402,0.010304374,0.32894838,0.009929693,0.00006530815,0.00035348075,0.00020846972,0.6446155],"genre_scores_gemma":[0.08004962,0.00867449,0.01451142,0.1955141,0.009601826,0.00044833258,0.0012079327,0.0006003748,0.689392],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9947747,0.0014464423,0.00015587588,0.0008680484,0.0020198848,0.00073506124],"domain_scores_gemma":[0.99014926,0.002785181,0.0003480925,0.0011273379,0.0036056603,0.0019844754],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007952679,0.0006163369,0.0005288523,0.0010449651,0.0039900965,0.010890168,0.002841284,0.005131406,0.133533],"category_scores_gemma":[0.021427592,0.00027775514,0.0005238536,0.0010689776,0.005091456,0.013902007,0.0076563656,0.007902512,0.047508273],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019438769,0.000040129642,0.0002627196,0.00009624095,0.0000048209768,0.00009115969,0.0005668413,0.00013003605,0.00013733076,0.59338784,0.34251413,0.06274921],"study_design_scores_gemma":[0.0000044801795,0.0000075220614,0.00011307334,0.00014392995,0.0000019338988,0.00010503121,0.0011854462,0.00012466035,0.00006491597,0.1764223,0.8218202,0.00000651651],"about_ca_topic_score_codex":0.0052491124,"about_ca_topic_score_gemma":0.0070827473,"teacher_disagreement_score":0.133533,"about_ca_system_score_codex":0.0039002057,"about_ca_system_score_gemma":0.010219221,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4413491739","doi":"10.1007/978-3-031-97788-6_13","title":"Expressions","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science","score_opus":0.010968159890456596,"score_gpt":0.30907673916540257,"score_spread":0.298108579274946,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491739","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0037639968,0.0008942143,0.19456069,0.0031834685,0.002140076,0.000477566,0.037467774,0.024699204,0.73281294],"genre_scores_gemma":[0.042769082,0.0018999286,0.10279234,0.002225007,0.00078647886,0.0007026184,0.06318922,0.017477259,0.7681581],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.997962,0.0003749412,0.00014480638,0.0006915894,0.00063962844,0.00018702791],"domain_scores_gemma":[0.99910516,0.00021331706,0.000039346163,0.00022222922,0.00037702805,0.00004292984],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011115934,0.0015867165,0.00063998933,0.002217458,0.0014886368,0.0044320184,0.0015392079,0.0009890797,0.2053221],"category_scores_gemma":[0.0033230945,0.00070705183,0.00085828005,0.0021807295,0.00097279187,0.0068679885,0.0035004034,0.0019865544,0.16739078],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014502398,0.00004947371,0.00044190622,0.00036835286,0.000020478306,0.00016522304,0.0006911439,0.00036855033,0.00409206,0.37270263,0.44346055,0.17749462],"study_design_scores_gemma":[0.000006914892,0.0000058917526,0.000117448435,0.000045954097,0.000009108162,0.000104700965,0.000119633056,0.0004462772,0.002155083,0.034326043,0.9626494,0.000013622471],"about_ca_topic_score_codex":0.003949989,"about_ca_topic_score_gemma":0.0033661888,"teacher_disagreement_score":0.2053221,"about_ca_system_score_codex":0.0016227523,"about_ca_system_score_gemma":0.0014909329,"threshold_uncertainty_score":0.6868708},"labels":[],"label_agreement":null},{"id":"W4413491740","doi":"10.1007/978-3-031-97788-6_14","title":"Queries","year":2025,"lang":"en","type":"book-chapter","venue":"Health information technology standards","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Computer science; Geography","score_opus":0.00902756598863077,"score_gpt":0.2962740940974728,"score_spread":0.28724652810884205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413491740","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007119205,0.0022342184,0.14195429,0.010339945,0.0023524798,0.001404752,0.1997,0.047594257,0.5873009],"genre_scores_gemma":[0.06295114,0.0033234672,0.071705915,0.0073927343,0.0011339785,0.0013147664,0.27517655,0.021493234,0.55550814],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980118,0.0003564809,0.00014212851,0.0004947041,0.0007858066,0.0002090232],"domain_scores_gemma":[0.9985024,0.00052221515,0.000051383297,0.00039248378,0.00044330835,0.00008830797],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012383349,0.0015995374,0.0009132931,0.002363678,0.00093952526,0.0046356525,0.0017631239,0.0016863943,0.26742166],"category_scores_gemma":[0.0064550284,0.0005554993,0.0009762094,0.0024728265,0.0006811786,0.0064688534,0.003487192,0.0015436257,0.15683793],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00027180483,0.00006670323,0.00067741604,0.0004773286,0.000028140732,0.00015572843,0.00035153,0.0006602298,0.0021137157,0.08348702,0.82918423,0.08252631],"study_design_scores_gemma":[0.00003435565,0.00001326356,0.00024423134,0.00006606619,0.000012267326,0.00010148332,0.00018133192,0.0016237965,0.0014826086,0.021429777,0.97479296,0.000017886616],"about_ca_topic_score_codex":0.009202159,"about_ca_topic_score_gemma":0.007478751,"teacher_disagreement_score":0.26742166,"about_ca_system_score_codex":0.0018148792,"about_ca_system_score_gemma":0.001438536,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W4413634702","doi":"10.64628/aam.nr7k6f4eu","title":"Why we should stop using acronyms like BIPOC","year":2023,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Computer science","score_opus":0.08691467059957698,"score_gpt":0.34162007372904907,"score_spread":0.2547054031294721,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413634702","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029842226,0.016875375,0.42269635,0.38943508,0.046696544,0.00053589133,0.0070825377,0.0253822,0.06145379],"genre_scores_gemma":[0.16299714,0.0068190307,0.6204398,0.13152784,0.0097675035,0.00069023995,0.008016546,0.019962184,0.039779708],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9832113,0.0067539676,0.0027518135,0.001900637,0.0044253194,0.00095703057],"domain_scores_gemma":[0.8930988,0.039248038,0.0077350354,0.010558423,0.045379143,0.003980574],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01370144,0.0013912334,0.0016678474,0.0064150756,0.00459659,0.007844339,0.0023331596,0.0032152857,0.018930811],"category_scores_gemma":[0.1021554,0.0009371418,0.0014105929,0.0064285076,0.004640809,0.021938194,0.004238632,0.00813691,0.021245308],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004248404,0.00012547556,0.005892604,0.0022094443,0.00014647281,0.0007404506,0.0035935058,0.0004583966,0.012708358,0.08044031,0.63211733,0.26114276],"study_design_scores_gemma":[0.00004467435,0.00004894544,0.0022560498,0.0011228389,0.00010386394,0.0011897712,0.0034237013,0.0020861998,0.0063400296,0.06913197,0.914097,0.00015496122],"about_ca_topic_score_codex":0.006395874,"about_ca_topic_score_gemma":0.006719914,"teacher_disagreement_score":0.018930811,"about_ca_system_score_codex":0.001997052,"about_ca_system_score_gemma":0.0043072603,"threshold_uncertainty_score":0.07246101},"labels":[],"label_agreement":null},{"id":"W4413971852","doi":"10.1177/10538135251365102","title":"Building Bridges: Establishing a Multiple Sclerosis Rehabilitation Research and Clinical Knowledge Mobilization Strategy","year":2025,"lang":"en","type":"article","venue":"Neurorehabilitation","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Manitoba; University Health Network; University of Toronto; Dalhousie University; Toronto Rehabilitation Institute; University of Alberta; University of Saskatchewan","funders":"Saskatchewan Health Research Foundation","keywords":"Multiple sclerosis; Mobilization; Rehabilitation; Physical medicine and rehabilitation; Medicine; Business; Physical therapy; Political science; Psychiatry","score_opus":0.10893333106307751,"score_gpt":0.42249051359579887,"score_spread":0.31355718253272136,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4413971852","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.107958116,0.006378208,0.28846475,0.39690465,0.0041330806,0.059745234,0.00060377,0.0019293233,0.13388291],"genre_scores_gemma":[0.29822573,0.0030715785,0.63416183,0.023528442,0.00047246733,0.022476275,0.000718166,0.00025623845,0.0170893],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.8897389,0.07010547,0.0062605925,0.006189251,0.014208854,0.013496938],"domain_scores_gemma":[0.8440439,0.044976726,0.007948341,0.009095475,0.039902512,0.0540331],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.23540267,0.0017392411,0.0014119531,0.010634173,0.022534367,0.022467805,0.010246464,0.011309546,0.009493396],"category_scores_gemma":[0.11788815,0.0017571419,0.0022794183,0.0050157616,0.012551083,0.019053048,0.051297642,0.0111295255,0.0034553818],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032186694,0.0027884475,0.01189216,0.0044589825,0.000199091,0.002697821,0.23761824,0.0023352196,0.0047320314,0.08756196,0.07607188,0.5693223],"study_design_scores_gemma":[0.0006490141,0.001304821,0.012791421,0.013303342,0.00029482433,0.0009956049,0.2900764,0.0064997673,0.0037025623,0.118872784,0.550977,0.0005324545],"about_ca_topic_score_codex":0.04302501,"about_ca_topic_score_gemma":0.09169876,"teacher_disagreement_score":0.23540267,"about_ca_system_score_codex":0.0520679,"about_ca_system_score_gemma":0.37513405,"threshold_uncertainty_score":0.9428846},"labels":[],"label_agreement":null},{"id":"W4414015785","doi":"10.11159/icbes25.120","title":"ACUITEE: A Comprehensive Tool for Visualization, Editing and Curating textual Annotations in Clinical Data","year":2025,"lang":"en","type":"article","venue":"Proceedings of the World Congress on Electrical Engineering and Computer Systems and Science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Computer science; Visualization; Data visualization; Information retrieval; Information visualization; World Wide Web; Human–computer interaction; Natural language processing; Artificial intelligence","score_opus":0.021804540202201714,"score_gpt":0.3175173256400033,"score_spread":0.29571278543780155,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414015785","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002956401,0.0015681438,0.5123538,0.001554069,0.0004464726,0.0012387979,0.057727594,0.4119913,0.010163399],"genre_scores_gemma":[0.02678122,0.0016924608,0.82227963,0.0019619395,0.00032214593,0.0032061457,0.08796948,0.043135963,0.012651013],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953501,0.0014895448,0.0007643044,0.0007556893,0.0014571509,0.00018314901],"domain_scores_gemma":[0.9785396,0.014329995,0.0012149073,0.0025777367,0.002545353,0.00079244986],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007492365,0.0026734956,0.0012194883,0.008132225,0.0014031363,0.0047467337,0.0030890633,0.0020570348,0.04769832],"category_scores_gemma":[0.03001755,0.0013422231,0.0018181942,0.005218509,0.0008391149,0.004799545,0.007004041,0.0023064634,0.022259427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007722207,0.00014312041,0.0022250363,0.004572855,0.00032937556,0.0011963228,0.0026138918,0.0018828394,0.016792808,0.009922445,0.61861825,0.34093082],"study_design_scores_gemma":[0.0003046207,0.00013322,0.0049051642,0.0012672258,0.0001446266,0.0021961564,0.00058985816,0.026492441,0.022998793,0.02004082,0.9205732,0.0003538859],"about_ca_topic_score_codex":0.0046036863,"about_ca_topic_score_gemma":0.009004837,"teacher_disagreement_score":0.04769832,"about_ca_system_score_codex":0.0012935633,"about_ca_system_score_gemma":0.0038855595,"threshold_uncertainty_score":0.15956676},"labels":[],"label_agreement":null},{"id":"W4414041908","doi":"10.2196/77837","title":"Large Language Model–Enhanced Drug Repositioning Knowledge Extraction via Long Chain-of-Thought: Development and Evaluation Study","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Drug repositioning; Adaptability; Knowledge base; Knowledge extraction; Knowledge acquisition; Information extraction; Domain knowledge","score_opus":0.014303347884581154,"score_gpt":0.3586234231411691,"score_spread":0.3443200752565879,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414041908","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.26254705,0.0111537995,0.6616972,0.0021256935,0.0005479816,0.0014150176,0.008776669,0.04429173,0.0074449196],"genre_scores_gemma":[0.55694443,0.0027287023,0.41221327,0.0008251656,0.00010562417,0.0005377811,0.023141855,0.00041124914,0.0030919318],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9983499,0.00069516053,0.0001680704,0.0003811758,0.0003102671,0.0000955191],"domain_scores_gemma":[0.9955929,0.002986123,0.00018417061,0.0004734209,0.00063127035,0.00013212701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035536285,0.0014167427,0.0011002419,0.0018225505,0.00040150725,0.0009655866,0.0019367572,0.0010803059,0.0031537947],"category_scores_gemma":[0.0070607495,0.00041372175,0.0017797566,0.0014771551,0.000432775,0.0023296238,0.0013127533,0.0014524458,0.0011472966],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012190065,0.001409006,0.005704121,0.0021781249,0.00077574747,0.0007573203,0.0002946015,0.21240468,0.014927983,0.0041837203,0.022091893,0.7340538],"study_design_scores_gemma":[0.00013485624,0.00028233527,0.0009553392,0.0000492621,0.00015269604,0.00014934585,0.00006592928,0.98484623,0.0076775854,0.0017382108,0.003919022,0.00002928134],"about_ca_topic_score_codex":0.012895951,"about_ca_topic_score_gemma":0.014437888,"teacher_disagreement_score":0.012895951,"about_ca_system_score_codex":0.0011209239,"about_ca_system_score_gemma":0.0026229946,"threshold_uncertainty_score":0.02564174},"labels":[],"label_agreement":null},{"id":"W4414117023","doi":"10.20944/preprints202509.0974.v1","title":"Integrating Unstructured EHR Data Using an FHIR-Based System: A Case Study with Problems List Data and FHIR IPS Model","year":2025,"lang":"en","type":"preprint","venue":"Preprints.org","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"École de Technologie Supérieure","funders":"","keywords":"Interoperability; Pipeline (software); Semantic interoperability; Key (lock); Resource (disambiguation); Component (thermodynamics); Unstructured data; Data model (GIS); Natural language","score_opus":0.29113878750762623,"score_gpt":0.4100422322737445,"score_spread":0.11890344476611825,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414117023","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8748315,0.0006974859,0.10054197,0.005434786,0.000116921205,0.0014357964,0.005326454,0.0022885045,0.009326565],"genre_scores_gemma":[0.8198148,0.00058269023,0.16851526,0.0007789925,0.000057452784,0.000435749,0.0056482656,0.00034271515,0.00382405],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9924901,0.0040424867,0.000750236,0.000760963,0.0016089365,0.0003472386],"domain_scores_gemma":[0.9840911,0.01120589,0.000745527,0.0016367045,0.0018024504,0.000518378],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007616083,0.00058774266,0.0005846641,0.0025878956,0.001881131,0.0032957725,0.0016976584,0.0029770886,0.001395016],"category_scores_gemma":[0.018697046,0.0003809077,0.0010839903,0.003869334,0.0013758321,0.0033906547,0.0021381371,0.0013516968,0.00070343324],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0025282074,0.0050107753,0.21405058,0.003680559,0.0005366158,0.17945643,0.07018135,0.060575187,0.03256448,0.033690587,0.0348219,0.36290345],"study_design_scores_gemma":[0.00066439837,0.0019147415,0.109162346,0.0011732166,0.00072686066,0.056433003,0.07300066,0.4099593,0.081842326,0.02084187,0.24370277,0.0005786534],"about_ca_topic_score_codex":0.023735674,"about_ca_topic_score_gemma":0.024776807,"teacher_disagreement_score":0.023735674,"about_ca_system_score_codex":0.0030737035,"about_ca_system_score_gemma":0.0021612716,"threshold_uncertainty_score":0.047195017},"labels":[],"label_agreement":null},{"id":"W4414120531","doi":"10.7554/elife.94909.2","title":"Mining the neuroimaging literature","year":2025,"lang":"en","type":"article","venue":"eLife","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University; Montreal Neurological Institute and Hospital","funders":"National Institute of Mental Health; Canadian Institutes of Health Research; National Institutes of Health; Natural Sciences and Engineering Research Council of Canada; Fondation Brain Canada; Fonds de recherche du Québec; Canada First Research Excellence Fund; Chan Zuckerberg Initiative; Fonds de Recherche du Québec - Santé; Michael J. Fox Foundation for Parkinson's Research","keywords":"Workflow; Upload; Metadata; Information extraction; Biomedical text mining; Task (project management); Source code","score_opus":0.008659507707274094,"score_gpt":0.2769466966464721,"score_spread":0.268287188939198,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414120531","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.05935209,0.12581614,0.30572054,0.02816493,0.004420366,0.0040996373,0.29702082,0.011682405,0.16372305],"genre_scores_gemma":[0.115369365,0.0719631,0.5275864,0.0050454433,0.0030750777,0.004170818,0.24993591,0.00267655,0.02017728],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99381363,0.0014943135,0.001251625,0.0012987829,0.0019161784,0.00022542199],"domain_scores_gemma":[0.9741911,0.012041475,0.0026265406,0.0021888872,0.008195198,0.0007566713],"candidate_categories":["bibliometrics"],"consensus_categories":[],"category_scores_codex":[0.0074823443,0.0011721562,0.0011536737,0.058212698,0.001919454,0.0042365394,0.002355968,0.0013351967,0.0121057285],"category_scores_gemma":[0.041342236,0.00061621726,0.0017241177,0.026582371,0.0011212635,0.0037229252,0.0042237607,0.0013740767,0.0079957405],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00020186436,0.000119106335,0.0090379985,0.015065179,0.00053906627,0.003995533,0.0031454256,0.0012977339,0.01031735,0.024792839,0.20853783,0.72295004],"study_design_scores_gemma":[0.000054400985,0.000053130705,0.009457384,0.0054542343,0.0006842574,0.0023301856,0.0024478238,0.0034437957,0.006891823,0.04625491,0.9228408,0.000087222186],"about_ca_topic_score_codex":0.0067990003,"about_ca_topic_score_gemma":0.013613829,"teacher_disagreement_score":0.9417873,"about_ca_system_score_codex":0.0019688443,"about_ca_system_score_gemma":0.009297074,"threshold_uncertainty_score":0.04049772},"labels":[],"label_agreement":null},{"id":"W4414184600","doi":"10.2196/76776","title":"Improving the Reporting Quality of Studies on Information Extraction From Clinical Texts: Protocol for the Development of a Consensus-Based Reporting Guideline","year":2025,"lang":"en","type":"article","venue":"JMIR Research Protocols","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Guideline; Protocol (science); Quality (philosophy); Data extraction; Information quality; MEDLINE; Quality management","score_opus":0.583009072102184,"score_gpt":0.6836492231874169,"score_spread":0.10064015108523294,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414184600","genre_codex":"protocol","genre_gemma":"protocol","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"protocol","genre_consensus":"protocol","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00117272,0.0011083476,0.027091986,0.003893309,0.0011938801,0.9574759,0.0048357034,0.0006821409,0.0025461584],"genre_scores_gemma":[0.001225594,0.00042911715,0.04376994,0.00070149056,0.00006680653,0.9523846,0.0010446742,0.0000636869,0.00031401886],"study_design_codex":"systematic_review","study_design_gemma":"not_applicable","domain_scores_codex":[0.41567427,0.32891813,0.20436694,0.014015492,0.029354444,0.0076707276],"domain_scores_gemma":[0.2975265,0.33351257,0.09311701,0.06614772,0.20128044,0.008415682],"candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.4646188,0.0056641046,0.010085805,0.019390041,0.0066658715,0.014033345,0.008902637,0.015324556,0.024979662],"category_scores_gemma":[0.5794419,0.0069862874,0.018790271,0.018433953,0.008949353,0.011559403,0.012775753,0.016373055,0.012856731],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006798369,0.0012960866,0.004634458,0.38726452,0.0042933035,0.0012852475,0.0299339,0.004902553,0.005669491,0.038279213,0.27214688,0.24349602],"study_design_scores_gemma":[0.020045236,0.0018325314,0.010969947,0.42547235,0.0038631193,0.0008345372,0.009295479,0.008428374,0.009045529,0.041481927,0.46746406,0.0012668492],"about_ca_topic_score_codex":0.0049291463,"about_ca_topic_score_gemma":0.005576082,"teacher_disagreement_score":0.5353812,"about_ca_system_score_codex":0.024657302,"about_ca_system_score_gemma":0.09005623,"threshold_uncertainty_score":0.6602204},"labels":[],"label_agreement":null},{"id":"W4414189579","doi":"10.1136/bmjdhai-2025-000014","title":"Optimising large language models for clinical information extraction: a benchmarking study in the context of ulcerative colitis research","year":2025,"lang":"en","type":"article","venue":"BMJ Digital Health & AI","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Artificial Intelligence in Medicine (Canada)","funders":"Clinical and Translational Science Institute, University of California, Los Angeles; Clinical and Translational Science Institute, University of California, San Francisco","keywords":"Benchmarking; Context (archaeology); Adaptation (eye); Oracle; Language model; Set (abstract data type); Predictive modelling; Colonoscopy","score_opus":0.11497063511652149,"score_gpt":0.5230203048782074,"score_spread":0.4080496697616859,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414189579","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6816127,0.0058194315,0.27555737,0.0033223245,0.000408588,0.0011934413,0.0073755123,0.019818954,0.0048918],"genre_scores_gemma":[0.771923,0.00095693406,0.20784768,0.00059015217,0.00010644674,0.00067125104,0.015409132,0.0009125505,0.0015828643],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9874959,0.009346471,0.00077206915,0.0015501648,0.0006260784,0.00020926427],"domain_scores_gemma":[0.9216779,0.069966815,0.0011412716,0.0037952878,0.0029835757,0.0004351848],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.01933746,0.0016412617,0.0008698837,0.0017141904,0.00050755456,0.0020250229,0.0022567336,0.0016492546,0.0020712023],"category_scores_gemma":[0.06170137,0.0007287963,0.0016757963,0.0020004078,0.0008421748,0.0032311894,0.0019405187,0.002243378,0.0011071782],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0028230567,0.0012269986,0.024489934,0.0024195223,0.0011242471,0.0007399291,0.0017280657,0.5579401,0.0059963586,0.0029923748,0.016833343,0.38168603],"study_design_scores_gemma":[0.00036909254,0.00063503097,0.0053952523,0.0001491724,0.00030797484,0.00025384035,0.00044095886,0.97656095,0.00591034,0.0043341755,0.0055528916,0.000090264664],"about_ca_topic_score_codex":0.012079808,"about_ca_topic_score_gemma":0.012775487,"teacher_disagreement_score":0.9806625,"about_ca_system_score_codex":0.0023218587,"about_ca_system_score_gemma":0.0020259665,"threshold_uncertainty_score":0.102267504},"labels":[],"label_agreement":null},{"id":"W4414192323","doi":"10.1101/2025.09.10.675443","title":"Evaluating Language Models for Biomedical Fact-Checking: A Benchmark Dataset for Cancer Variant Interpretation Verification","year":2025,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"National Cancer Institute; Canadian Institutes of Health Research; National Institutes of Health; University of British Columbia","keywords":"Pipeline (software); Benchmark (surveying); Interpretation (philosophy); Process (computing); Triage; Path (computing)","score_opus":0.041604507087200246,"score_gpt":0.3429288617450822,"score_spread":0.30132435465788193,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414192323","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23191296,0.0076464685,0.091474615,0.010226609,0.00137191,0.0020445036,0.5937293,0.04025989,0.021333717],"genre_scores_gemma":[0.13971657,0.0008530078,0.1333258,0.0015649883,0.00016859292,0.0007870006,0.71835434,0.0012188669,0.0040108273],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9941392,0.0018032056,0.0008583506,0.0014857738,0.0014679992,0.00024550938],"domain_scores_gemma":[0.96572524,0.021752436,0.0015288166,0.0047196704,0.0054895002,0.0007842739],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0077910493,0.0014766811,0.000618839,0.0049184063,0.0015034425,0.002251604,0.0028629927,0.0029540858,0.0047469093],"category_scores_gemma":[0.040375322,0.00044067157,0.001621647,0.0033826143,0.001195253,0.0025683849,0.002605587,0.0018084937,0.0032014481],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019980785,0.001245258,0.037798952,0.005447359,0.0006796036,0.0030140507,0.0013243502,0.041840084,0.0147894705,0.015767511,0.6934569,0.18263839],"study_design_scores_gemma":[0.0018281484,0.0006886052,0.033249162,0.0011313547,0.00054257736,0.0039485428,0.0019517731,0.28073594,0.047312494,0.035821185,0.59247446,0.00031570063],"about_ca_topic_score_codex":0.017860152,"about_ca_topic_score_gemma":0.028762966,"teacher_disagreement_score":0.99220896,"about_ca_system_score_codex":0.002510144,"about_ca_system_score_gemma":0.0043200585,"threshold_uncertainty_score":0.04120344},"labels":[],"label_agreement":null},{"id":"W4414536046","doi":"10.21203/rs.3.rs-7511758/v2","title":"Spatial Reasoning AI in Clinical Workflows: A Scoping Review of Translational Applications","year":2025,"lang":"en","type":"preprint","venue":"Research Square","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Columbia College","funders":"","keywords":"Workflow; Segmentation; Identification (biology); Precision medicine; Class (philosophy); Clinical trial","score_opus":0.09367009197818249,"score_gpt":0.5122688715109567,"score_spread":0.4185987795327742,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414536046","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00038320658,0.9935568,0.0030768032,0.0016318283,0.00012965046,0.00003904212,0.00009153055,0.00002965526,0.0010614247],"genre_scores_gemma":[0.008530467,0.98004675,0.009425478,0.0009828878,0.00035087077,0.00010198786,0.0002610361,0.000033361895,0.00026721752],"study_design_codex":"design_other","study_design_gemma":"systematic_review","domain_scores_codex":[0.99198884,0.0032449157,0.0017842832,0.00095251936,0.0017960094,0.00023326179],"domain_scores_gemma":[0.86260617,0.12382364,0.0028762503,0.00269797,0.0074982475,0.00049772044],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.028869592,0.0012355596,0.0027693047,0.013735769,0.0010005597,0.0068744086,0.0033443864,0.002943644,0.003345229],"category_scores_gemma":[0.064990416,0.0010861249,0.0034219916,0.015870523,0.0033229457,0.008538556,0.0031043065,0.0027825334,0.0008819809],"study_design_candidate":"systematic_review","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024367085,0.00013354492,0.0012746975,0.1225044,0.0013496206,0.000096065494,0.00086778717,0.0027633756,0.0006616715,0.0227247,0.013061481,0.83431894],"study_design_scores_gemma":[0.00013818435,0.00029191704,0.0040983516,0.30523193,0.005482969,0.00095630187,0.0020361794,0.0074520065,0.0026633637,0.07880381,0.5926771,0.00016786325],"about_ca_topic_score_codex":0.008869837,"about_ca_topic_score_gemma":0.008963702,"teacher_disagreement_score":0.028869592,"about_ca_system_score_codex":0.004031212,"about_ca_system_score_gemma":0.0119129615,"threshold_uncertainty_score":0.15267879},"labels":[],"label_agreement":null},{"id":"W4414598938","doi":"10.1101/2025.09.10.25334730","title":"EAGLE-AI: A large language model workflow for automated extraction and scoring of literature evidence linking genes to autism spectrum disorder","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; SickKids Foundation; Hospital for Sick Children","funders":"University of Toronto; Hospital for Sick Children; Autism Speaks","keywords":"Workflow; Autism; Autism spectrum disorder; Set (abstract data type); Data curation; Component (thermodynamics); Data set; Genomics","score_opus":0.020995557040339514,"score_gpt":0.3390197854823722,"score_spread":0.3180242284420327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414598938","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.009843997,0.0018036024,0.5914395,0.0022558603,0.00031685515,0.0011075553,0.12929346,0.25738612,0.00655299],"genre_scores_gemma":[0.039928515,0.00079841446,0.827262,0.00070172886,0.00012460323,0.0009008757,0.11728736,0.008208789,0.004787701],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9964839,0.00095424306,0.0006577226,0.0008002055,0.0009576417,0.00014618518],"domain_scores_gemma":[0.98745173,0.0069194706,0.001011398,0.0018502494,0.0022056338,0.00056148507],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00980639,0.0018947879,0.0011482789,0.009329503,0.0011544351,0.003944427,0.0020980851,0.0011256406,0.017639074],"category_scores_gemma":[0.02194284,0.0009215701,0.0024809805,0.0035203358,0.0004968141,0.00238584,0.0034088113,0.0014325831,0.011532622],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013280077,0.000269501,0.009967417,0.009042396,0.0015345509,0.0016721041,0.0016957646,0.011717118,0.06623996,0.022041896,0.47661403,0.39787722],"study_design_scores_gemma":[0.00063622865,0.00029126115,0.010089849,0.0012331579,0.00068758905,0.00184361,0.000981123,0.17190908,0.10242105,0.06327002,0.64616007,0.00047704746],"about_ca_topic_score_codex":0.0052523497,"about_ca_topic_score_gemma":0.014127116,"teacher_disagreement_score":0.017639074,"about_ca_system_score_codex":0.0012747343,"about_ca_system_score_gemma":0.0059779976,"threshold_uncertainty_score":0.0590086},"labels":[],"label_agreement":null},{"id":"W4414620780","doi":"10.37044/osf.io/wj8bz_v1","title":"Translating and Formalizing the MIRAGE Guidelines to a Prototype MIRAGE Ontology and DCAT3 Extension Vocabulary for Glycomics Data Management","year":2025,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Bioscience Database Center; Canadian Glycomics Network; Japan Science and Technology Agency; Ministry of Education, Culture, Sports, Science and Technology","keywords":"Glycomics; Ontology; Semantic Web; Metadata; RDF; Standardization; Controlled vocabulary; Vocabulary; Interoperability","score_opus":0.12344301981468878,"score_gpt":0.39323873865715925,"score_spread":0.26979571884247044,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414620780","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0066744047,0.00021163082,0.9523063,0.0013576348,0.0002541329,0.0011829955,0.008493906,0.019246792,0.010272115],"genre_scores_gemma":[0.048430398,0.00072166877,0.8862507,0.0014777083,0.00012669674,0.0015748893,0.046573684,0.0067342315,0.008109987],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99346465,0.0011060835,0.0014868161,0.001002596,0.002392063,0.0005478215],"domain_scores_gemma":[0.992465,0.0015826869,0.0006144623,0.002415233,0.0026303795,0.00029215208],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010014521,0.0012456884,0.00080204307,0.00309781,0.0016273645,0.0057970136,0.0028093597,0.0017322848,0.005451632],"category_scores_gemma":[0.010912117,0.0012059152,0.0026170763,0.0024091615,0.002025901,0.007161751,0.005225192,0.0040328437,0.0034721557],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00046652323,0.00083488243,0.005843981,0.0018149412,0.00020963445,0.0014046751,0.004024283,0.021751225,0.046939235,0.5577143,0.13263276,0.22636357],"study_design_scores_gemma":[0.00011408562,0.00007190063,0.0014816268,0.000919763,0.000110530455,0.00063284417,0.0007416716,0.06649884,0.043073013,0.09974209,0.78642595,0.00018774628],"about_ca_topic_score_codex":0.027603101,"about_ca_topic_score_gemma":0.03247184,"teacher_disagreement_score":0.98998547,"about_ca_system_score_codex":0.003453842,"about_ca_system_score_gemma":0.010727636,"threshold_uncertainty_score":0.05488485},"labels":[],"label_agreement":null},{"id":"W4414755636","doi":"10.1038/s41586-025-09648-x","title":"Publisher Correction: Multimodal cell maps as a foundation for structural and functional genomics","year":2025,"lang":"en","type":"erratum","venue":"Nature","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canada's Michael Smith Genome Sciences Centre; University of British Columbia","funders":"","keywords":"Foundation (evidence); Functional genomics; Genomics; Genome Biology; Functional analysis","score_opus":0.00695335879531028,"score_gpt":0.25834571231670483,"score_spread":0.25139235352139455,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414755636","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0003162586,0.0012056719,0.014852189,0.04243508,0.91711634,0.00004079665,0.0044380063,0.0025391201,0.017056579],"genre_scores_gemma":[0.018254768,0.005259966,0.05723357,0.04211791,0.12648089,0.00036264322,0.010601399,0.011700384,0.7279884],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99498135,0.00078600604,0.0006969409,0.00069605635,0.0026393924,0.00020025548],"domain_scores_gemma":[0.96207243,0.009130156,0.0010966111,0.005163236,0.021608729,0.00092875067],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050433013,0.0021662011,0.0013525486,0.0042459015,0.0034320315,0.0053547584,0.0033253375,0.0039927294,0.09431117],"category_scores_gemma":[0.060025778,0.0012384681,0.0014563916,0.0033606295,0.0028280201,0.003909305,0.0025600025,0.010683783,0.057094567],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020237028,0.000004502011,0.00005267826,0.00010404892,0.000012047363,0.00013224962,0.00004394349,0.00009927082,0.00012234227,0.005706897,0.9840583,0.009643449],"study_design_scores_gemma":[0.000011601228,0.000006341762,0.00023397263,0.00011631011,0.00002498275,0.0003192362,0.000043512926,0.00036531093,0.0005898024,0.004322702,0.99394107,0.000025165584],"about_ca_topic_score_codex":0.012741528,"about_ca_topic_score_gemma":0.020472359,"teacher_disagreement_score":0.09431117,"about_ca_system_score_codex":0.0032571245,"about_ca_system_score_gemma":0.004930662,"threshold_uncertainty_score":0.31550235},"labels":[],"label_agreement":null},{"id":"W4414815947","doi":"","title":"Cartographie des initiatives en cours dans les agences sanitaires concernant l’utilisation des outils d’IA pour la revue de littérature","year":2025,"lang":"fr","type":"report","venue":"HAL (Le Centre pour la Communication Scientifique Directe)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.02547352259400315,"score_gpt":0.27955505758542404,"score_spread":0.2540815349914209,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414815947","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.057829145,0.16782252,0.08988764,0.06591282,0.0065822317,0.0016731741,0.016517378,0.004317556,0.5894575],"genre_scores_gemma":[0.31388226,0.23322095,0.13775511,0.010690559,0.0015697674,0.003177062,0.025923906,0.0023239949,0.27145636],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9921434,0.0030828884,0.00068815536,0.0008143115,0.0026859276,0.0005853162],"domain_scores_gemma":[0.97773963,0.010789257,0.0020647256,0.0014240773,0.006826991,0.0011552409],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.008335857,0.0007528009,0.0004416702,0.011379773,0.0021613056,0.008523733,0.0012328277,0.0016052168,0.031413775],"category_scores_gemma":[0.016021501,0.00054154405,0.0011369777,0.015933981,0.0021380556,0.008104157,0.0027487746,0.0027963386,0.00717285],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021265405,0.00011095793,0.011493322,0.008216325,0.000091932365,0.00063670374,0.021747606,0.002100905,0.0033677693,0.23539777,0.22228163,0.49434248],"study_design_scores_gemma":[0.0000071792024,0.00003899308,0.008977128,0.002512185,0.00002658088,0.00020930642,0.0046034465,0.00031321112,0.0006764305,0.004132637,0.97847044,0.000032417596],"about_ca_topic_score_codex":0.03599371,"about_ca_topic_score_gemma":0.0413837,"teacher_disagreement_score":0.9916642,"about_ca_system_score_codex":0.007988228,"about_ca_system_score_gemma":0.012692132,"threshold_uncertainty_score":0.105089486},"labels":[],"label_agreement":null},{"id":"W4414857549","doi":"10.1093/genetics/iyaf215","title":"Mondo: integrating disease terminology across communities","year":2025,"lang":"en","type":"article","venue":"Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":14,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Jewish General Hospital","funders":"Office of AIDS Research; Basic Energy Sciences; National Institute of Child Health and Human Development; National Institute of Neurological Disorders and Stroke; National Institute of Arthritis and Musculoskeletal and Skin Diseases; National Institute of Allergy and Infectious Diseases; National Cancer Institute; Center for Information Technology; School of Veterinary Science, University of Queensland; Office of Science; U.S. National Library of Medicine; Defense Advanced Research Projects Agency; Advanced Research Projects Agency; National Institutes of Health; National Human Genome Research Institute; Wellcome Trust; Division of Intramural Research, National Institute of Allergy and Infectious Diseases; Norges Idrettshøgskole; Eunice Kennedy Shriver National Institute of Child Health and Human Development; U.S. Department of Energy; University of California, San Francisco; Advanced Research Projects Agency for Health","keywords":"Interoperability; SNOMED CT; Terminology; Disease; Ontology; Clinical decision support system; Coding (social sciences); Inheritance (genetic algorithm); Decision support system; Precision medicine","score_opus":0.02006117502985651,"score_gpt":0.3255779963364236,"score_spread":0.3055168213065671,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414857549","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018328803,0.0058236565,0.867262,0.010056825,0.0015704766,0.0019455036,0.041056007,0.017702715,0.036254063],"genre_scores_gemma":[0.11621737,0.0037656229,0.7903262,0.0037976464,0.00056855706,0.0017112304,0.07418202,0.0025272039,0.0069041853],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.98813486,0.0040310957,0.0022573469,0.00234072,0.002576315,0.0006597357],"domain_scores_gemma":[0.9859432,0.0049543385,0.0015482042,0.0040262016,0.0021867566,0.0013413575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013736996,0.0010833492,0.0013819806,0.014804967,0.003237254,0.0079369685,0.0033809347,0.0027384483,0.0048434054],"category_scores_gemma":[0.041312028,0.00085100025,0.0035398973,0.012341697,0.0022795636,0.014864382,0.02452474,0.0022345367,0.002357963],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004527123,0.00012182681,0.012737472,0.0024593796,0.00035811943,0.0011171442,0.0074063484,0.005768692,0.003931453,0.5772413,0.113081045,0.27532446],"study_design_scores_gemma":[0.000092211194,0.000059607093,0.0038558987,0.0013125598,0.00018874696,0.0008361203,0.0018854322,0.016939027,0.0015534365,0.3272592,0.6458694,0.00014844109],"about_ca_topic_score_codex":0.02336385,"about_ca_topic_score_gemma":0.021121796,"teacher_disagreement_score":0.02336385,"about_ca_system_score_codex":0.0041146455,"about_ca_system_score_gemma":0.009526385,"threshold_uncertainty_score":0.072649},"labels":[],"label_agreement":null},{"id":"W4414886508","doi":"10.59350/fa3qb-76m68","title":"Montreal BioJava Bootcamp Announced","year":2003,"lang":"en","type":"preprint","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Sequence (biology); Focus (optics); Context (archaeology); Editorial board","score_opus":0.01773147944329089,"score_gpt":0.2728545852523451,"score_spread":0.25512310580905423,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4414886508","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0036802893,0.009119755,0.0059619183,0.05673072,0.024681062,0.00048921414,0.021508547,0.0064460137,0.87138253],"genre_scores_gemma":[0.004419589,0.0011650297,0.0019153828,0.0016786918,0.0006670307,0.000069953654,0.004895497,0.00069315446,0.9844956],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99823403,0.00011003572,0.00002129286,0.00032772182,0.0009276476,0.0003793259],"domain_scores_gemma":[0.9961422,0.00017228689,0.00006347365,0.00032097643,0.0018496108,0.00145143],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0023972222,0.0012933458,0.00082023407,0.0016917225,0.006286312,0.0067901476,0.0016373192,0.0023411813,0.4520676],"category_scores_gemma":[0.0033554477,0.00072993967,0.0008095186,0.0018627021,0.0009988877,0.0021242443,0.0026697323,0.0033474981,0.17376928],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040015162,0.0000133489175,0.000108740955,0.0000127284575,0.0000018562614,0.000067486486,0.000026690668,0.000030974443,0.000359851,0.004002925,0.9851841,0.010151279],"study_design_scores_gemma":[0.00000646618,0.000005561648,0.00025028302,0.000007843518,8.491422e-7,0.000017737748,0.000023796041,0.000044548022,0.00013918635,0.00032564145,0.9991736,0.000004457256],"about_ca_topic_score_codex":0.5074538,"about_ca_topic_score_gemma":0.7233091,"teacher_disagreement_score":0.5074538,"about_ca_system_score_codex":0.011720129,"about_ca_system_score_gemma":0.016964281,"threshold_uncertainty_score":0.99089384},"labels":[],"label_agreement":null},{"id":"W4415179083","doi":"10.1109/icjece.2025.3607372","title":"SAACT: Semiautomated Annotation of Computerized Tomography Data","year":2025,"lang":"en","type":"article","venue":"Canadian Journal of Electrical and Computer Engineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"Ontario Tech University","funders":"","keywords":"Preprocessor; Annotation; Segmentation; Python (programming language); Domain (mathematical analysis); Noise (video); Volume (thermodynamics); Pattern recognition (psychology)","score_opus":0.007314087102822671,"score_gpt":0.2184932664125809,"score_spread":0.21117917930975824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415179083","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005833645,0.00012874533,0.9587912,0.00014256356,0.00005394151,0.0001352477,0.0019451557,0.031797115,0.0011724336],"genre_scores_gemma":[0.07638106,0.00028325233,0.9048785,0.00021579467,0.000048236336,0.0004883264,0.012883767,0.0021754391,0.0026456756],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982362,0.00031536882,0.0001561442,0.00053237984,0.00066059805,0.000099370336],"domain_scores_gemma":[0.99648273,0.00088422024,0.00039374267,0.0011914824,0.0009544093,0.000093444076],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022909467,0.0017223164,0.00095442927,0.0030531175,0.0007448938,0.0023533762,0.0026302326,0.0013536182,0.004280212],"category_scores_gemma":[0.005941351,0.00079788396,0.0016293544,0.0023356148,0.0010910399,0.0024222303,0.0031572373,0.0017539533,0.0037876861],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006261175,0.00030371314,0.005639614,0.0010128339,0.00029090594,0.0004898282,0.00071532966,0.06903567,0.10597282,0.0140424585,0.058211684,0.74365896],"study_design_scores_gemma":[0.000039521205,0.0001645745,0.004292968,0.000116501426,0.00006156378,0.00068080163,0.0002822334,0.818068,0.09417321,0.01918167,0.06283112,0.000107768916],"about_ca_topic_score_codex":0.0054280865,"about_ca_topic_score_gemma":0.008622686,"teacher_disagreement_score":0.0054280865,"about_ca_system_score_codex":0.0009756962,"about_ca_system_score_gemma":0.0024157693,"threshold_uncertainty_score":0.014318705},"labels":[],"label_agreement":null},{"id":"W4415216333","doi":"10.1101/2025.10.13.25337935","title":"A machine learning model to support the screening for methods guidance articles in MEDLINE: A performance evaluation of ASReview simulation mode","year":2025,"lang":"en","type":"preprint","venue":"medRxiv","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Support vector machine; Naive Bayes classifier; Artificial neural network; Feature (linguistics); Precision and recall; Search engine indexing; Feature extraction; Recall","score_opus":0.1593125589647672,"score_gpt":0.46397754032292265,"score_spread":0.30466498135815545,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415216333","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.81164604,0.0059155705,0.14666884,0.0028291831,0.0005751689,0.0015272206,0.008787811,0.01836098,0.003689097],"genre_scores_gemma":[0.88012826,0.0008558954,0.10745516,0.00064846006,0.00012312082,0.00063985056,0.008552903,0.00012273478,0.0014737116],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981805,0.00085574447,0.00023749312,0.000391822,0.00024633773,0.000088076165],"domain_scores_gemma":[0.97518754,0.019320324,0.001021686,0.00079320284,0.003175077,0.0005021651],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.007023941,0.001265827,0.0014467901,0.0034096388,0.00048788733,0.0014126941,0.0016435339,0.0014405827,0.0018362112],"category_scores_gemma":[0.023420751,0.00044993366,0.0013151247,0.001606378,0.00030594456,0.0014994007,0.0006587059,0.0010657279,0.0009844786],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005455406,0.0019670315,0.06608349,0.0018180413,0.0011775551,0.0004392669,0.000297772,0.61188394,0.006026681,0.0010534517,0.018213095,0.28558433],"study_design_scores_gemma":[0.00008213914,0.00021159681,0.001365113,0.00003823338,0.00007802726,0.000045503264,0.00002709387,0.9956664,0.0016385238,0.0003431108,0.00048961176,0.000014527303],"about_ca_topic_score_codex":0.01795795,"about_ca_topic_score_gemma":0.014845485,"teacher_disagreement_score":0.99297607,"about_ca_system_score_codex":0.0018190874,"about_ca_system_score_gemma":0.003423535,"threshold_uncertainty_score":0.037146628},"labels":[],"label_agreement":null},{"id":"W4415544440","doi":"10.1093/jamia/ocaf184","title":"An exploratory analysis of SNOMED CT national editions","year":2025,"lang":"en","type":"article","venue":"Journal of the American Medical Informatics Association","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Victoria","funders":"National Institutes of Health","keywords":"SNOMED CT; Consistency (knowledge bases); Exploratory analysis; Systematized Nomenclature of Medicine; Extension (predicate logic)","score_opus":0.008706236505678578,"score_gpt":0.30931860472644457,"score_spread":0.300612368220766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415544440","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8666426,0.0021141414,0.042084444,0.0013890584,0.00013264985,0.00207621,0.052471604,0.0014005441,0.031688713],"genre_scores_gemma":[0.84211135,0.0009467181,0.11046649,0.00044668117,0.00005787398,0.0025294533,0.03951261,0.0008718412,0.0030568962],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9850086,0.0062715947,0.0024373021,0.0015256567,0.004233189,0.00052363414],"domain_scores_gemma":[0.91949356,0.05157042,0.0064212536,0.0058188667,0.015998734,0.00069707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0200039,0.00047489817,0.00048854586,0.01626792,0.0013690209,0.0029065595,0.0009080111,0.00048261645,0.0040831883],"category_scores_gemma":[0.07783797,0.00035399903,0.0009872712,0.018645223,0.0015941915,0.0033135943,0.004057545,0.0007158775,0.0006391868],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022569178,0.00041174793,0.2964953,0.009513615,0.00066360505,0.0060542626,0.17436767,0.0054664416,0.02170954,0.042217646,0.038861,0.40198228],"study_design_scores_gemma":[0.00013503744,0.00058603013,0.43743524,0.005359256,0.0010175076,0.006974818,0.11904961,0.017817296,0.02341118,0.018792639,0.3689059,0.0005154914],"about_ca_topic_score_codex":0.008052312,"about_ca_topic_score_gemma":0.011849446,"teacher_disagreement_score":0.0200039,"about_ca_system_score_codex":0.003174191,"about_ca_system_score_gemma":0.0047339285,"threshold_uncertainty_score":0.105791986},"labels":[],"label_agreement":null},{"id":"W4415683348","doi":"10.1093/bioinformatics/btaf582","title":"Endowing protein language models with structural knowledge","year":2025,"lang":"en","type":"article","venue":"Bioinformatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Instituto de Ciencias del Mar y Limnología, Universidad Nacional Autónoma de México; Institute for Catastrophic Loss Reduction","keywords":"Code (set theory); Software; Source code; Knowledge representation and reasoning; Natural language; Language model","score_opus":0.01077066642079481,"score_gpt":0.26765891738078107,"score_spread":0.25688825095998624,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415683348","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035015192,0.0003559481,0.9567283,0.0008369753,0.000051300052,0.000048971582,0.0006458897,0.004429854,0.0018875727],"genre_scores_gemma":[0.52001756,0.0009182527,0.4671672,0.0005394681,0.00012596315,0.00028738735,0.0046061925,0.0007124434,0.005625514],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9994338,0.00017909845,0.000034409964,0.00015573883,0.0001555165,0.00004141358],"domain_scores_gemma":[0.9979042,0.0012017165,0.00016079961,0.0003901209,0.00026131188,0.00008175651],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011133994,0.0009349791,0.00050449435,0.00065983593,0.00029659012,0.0009826598,0.0015112638,0.00083956594,0.0031349906],"category_scores_gemma":[0.0072638057,0.0005082093,0.0007776124,0.00080158963,0.0006785199,0.0035405515,0.002305134,0.0020457434,0.002859422],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025055197,0.00023212547,0.0024347045,0.0004112076,0.00011079162,0.00019680212,0.000201916,0.691763,0.022148967,0.034581468,0.009201369,0.2384671],"study_design_scores_gemma":[0.000007658566,0.000021853371,0.00009505645,0.000009153392,0.000009446844,0.000025887946,0.000012588112,0.9751394,0.003494923,0.019714803,0.0014621943,0.0000070246647],"about_ca_topic_score_codex":0.0023843278,"about_ca_topic_score_gemma":0.004054407,"teacher_disagreement_score":0.0031349906,"about_ca_system_score_codex":0.00066057494,"about_ca_system_score_gemma":0.0010213417,"threshold_uncertainty_score":0.010487616},"labels":[],"label_agreement":null},{"id":"W4415816438","doi":"10.1093/genetics/iyaf237","title":"Xenbase: 25 years of integrating molecular and biomedical data from <i>Xenopus</i>","year":2025,"lang":"en","type":"article","venue":"Genetics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"National Institutes of Health","keywords":"Xenopus; Suite; Variety (cybernetics); Focus (optics); Translation (biology); Exome; Model organism; Genome; Software","score_opus":0.017829725943106056,"score_gpt":0.29629178783500937,"score_spread":0.2784620618919033,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415816438","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0060018282,0.009328503,0.09052444,0.0035914013,0.0011307723,0.00048358197,0.7691273,0.09063311,0.029179096],"genre_scores_gemma":[0.009306849,0.0049466216,0.060224008,0.0011896157,0.00014642565,0.00042076886,0.9107963,0.009124506,0.0038448242],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9969919,0.00048990134,0.00076345797,0.00055351434,0.0010225534,0.00017863089],"domain_scores_gemma":[0.9916157,0.0022590759,0.0009984446,0.002583623,0.0016771152,0.0008660221],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0051494585,0.0019293196,0.001715939,0.009999065,0.0012374842,0.0058487626,0.003695078,0.0018352179,0.030981656],"category_scores_gemma":[0.016423486,0.0012747991,0.0017121296,0.010877263,0.0011129328,0.00553258,0.0062445076,0.0022894966,0.029235762],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013863094,0.00013610668,0.005679577,0.007164112,0.0007039092,0.0009503296,0.00072834874,0.0032499672,0.018461067,0.019639084,0.7771296,0.16477159],"study_design_scores_gemma":[0.000085111715,0.00004881713,0.0034142344,0.0009100115,0.00019551971,0.00051706063,0.00014272492,0.0013962048,0.008443347,0.006829589,0.9779107,0.000106677886],"about_ca_topic_score_codex":0.010016123,"about_ca_topic_score_gemma":0.00978073,"teacher_disagreement_score":0.030981656,"about_ca_system_score_codex":0.0022192376,"about_ca_system_score_gemma":0.0052768267,"threshold_uncertainty_score":0.103643954},"labels":[],"label_agreement":null},{"id":"W4415991491","doi":"10.1108/978-1-83662-494-320251002","title":"Interactions – A CyberSystemic Model and an Observation Framework Proposal","year":2025,"lang":"en","type":"book-chapter","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Cybernet Systems Corporation (Canada)","funders":"","keywords":"Domain (mathematical analysis); Interaction model; Process (computing); Field (mathematics); Human interaction; Order (exchange)","score_opus":0.029779436928878175,"score_gpt":0.30679300522266234,"score_spread":0.27701356829378415,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4415991491","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010314687,0.0021760461,0.9343921,0.008803973,0.00024367456,0.00051404967,0.00064339046,0.0006887452,0.042223237],"genre_scores_gemma":[0.31895295,0.0038757492,0.6569036,0.0011524976,0.00039687465,0.001469964,0.0017073574,0.00022611108,0.015314888],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9945862,0.0016614244,0.00056230323,0.0012974859,0.0015029134,0.00038951647],"domain_scores_gemma":[0.9964371,0.0014468872,0.00040078492,0.00070234196,0.0007550715,0.0002578209],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0054996503,0.0010908705,0.00068103836,0.0033990054,0.0019479166,0.008025216,0.0030120825,0.0025075558,0.0042767623],"category_scores_gemma":[0.004849069,0.00090834516,0.002383128,0.0028303554,0.007450174,0.01602723,0.0052246,0.004540841,0.0010522382],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000015834927,0.000033761597,0.0008606333,0.00014024053,0.000022116408,0.0002031182,0.001992715,0.005857446,0.0005273235,0.9750963,0.0016603126,0.013590323],"study_design_scores_gemma":[0.00002012881,0.000101537065,0.0016957207,0.00045616925,0.00007038984,0.0005299001,0.0031812128,0.056386776,0.0010470102,0.8119768,0.124456756,0.000077618686],"about_ca_topic_score_codex":0.011765839,"about_ca_topic_score_gemma":0.00510579,"teacher_disagreement_score":0.011765839,"about_ca_system_score_codex":0.0042277193,"about_ca_system_score_gemma":0.0060156696,"threshold_uncertainty_score":0.030674338},"labels":[],"label_agreement":null},{"id":"W4416034984","doi":"10.18653/v1/2025.findings-emnlp.223","title":"Sequence Structure Aware Retriever for Procedural Document Retrieval: A New Dataset and Baseline","year":2025,"lang":"","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Fundamental Research Funds for the Central Universities; South China University of Technology; Natural Science Foundation of Guangxi Province; Natural Science Foundation of Guangdong Province; National Natural Science Foundation of China","keywords":"Baseline (sea); Sequence (biology); Labrador Retriever; Pattern recognition (psychology); Component (thermodynamics)","score_opus":0.021068560396122715,"score_gpt":0.33045302659210035,"score_spread":0.30938446619597765,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416034984","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16463587,0.011611872,0.05920341,0.0033018286,0.00094584445,0.0033683195,0.7031753,0.03800926,0.015748253],"genre_scores_gemma":[0.050510965,0.0011142971,0.093419634,0.0006792185,0.0001461641,0.0011794019,0.8476637,0.00048370665,0.004802971],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979234,0.00027278386,0.0002773889,0.00075271004,0.0005981911,0.00017552754],"domain_scores_gemma":[0.9964988,0.0009170397,0.00025926996,0.0011744659,0.0008783182,0.0002720074],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001819324,0.002135975,0.0013812673,0.006298163,0.0016840686,0.0020711392,0.0036326025,0.0032608388,0.0053707184],"category_scores_gemma":[0.0069771996,0.0003441679,0.0023010548,0.004814036,0.0010964527,0.0033265392,0.00201249,0.0023316103,0.008273577],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002272155,0.003741455,0.014619879,0.0061451504,0.00052222644,0.001468271,0.00043108445,0.01858342,0.037587147,0.004408333,0.5373257,0.37289524],"study_design_scores_gemma":[0.0016457145,0.00218317,0.050209187,0.00073730387,0.00069299695,0.0050391364,0.0018995117,0.14154796,0.061427318,0.011339794,0.72278094,0.00049687346],"about_ca_topic_score_codex":0.025231136,"about_ca_topic_score_gemma":0.037921112,"teacher_disagreement_score":0.025231136,"about_ca_system_score_codex":0.0021007243,"about_ca_system_score_gemma":0.0024414281,"threshold_uncertainty_score":0.050168514},"labels":[],"label_agreement":null},{"id":"W4416185268","doi":"10.1016/j.vhri.2025.101539","title":"Evaluating the Performance of Claude 3.7 Sonnet in Data Extraction Automation for Systematic Literature Reviews","year":2025,"lang":"en","type":"article","venue":"Value in Health Regional Issues","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"EVRAZ (Canada)","funders":"","keywords":"Systematic review; Sonnet; Data extraction; Automation; Extraction (chemistry)","score_opus":0.16782730204867372,"score_gpt":0.4849505993754512,"score_spread":0.31712329732677746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416185268","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.58485836,0.028418304,0.23739317,0.008917271,0.0013991984,0.004056791,0.0486543,0.06823418,0.018068392],"genre_scores_gemma":[0.3964772,0.0035688127,0.5573996,0.0007895121,0.00019876492,0.0015551193,0.03619536,0.0013057181,0.0025099702],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.97450644,0.013255631,0.0063842735,0.0023641104,0.002993556,0.00049593527],"domain_scores_gemma":[0.7296118,0.23737454,0.0077697765,0.0067907097,0.016710375,0.0017428403],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.046309438,0.001476104,0.0020354833,0.013926907,0.0013844458,0.004336926,0.0013541412,0.002468133,0.004571468],"category_scores_gemma":[0.17158957,0.0011154515,0.0030129177,0.008892041,0.00074196997,0.0040111523,0.0034992243,0.0010117963,0.0016566749],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.013114003,0.0015936136,0.19191022,0.045755092,0.012588038,0.0013627507,0.0045013023,0.078856945,0.01240839,0.022558028,0.064213276,0.55113834],"study_design_scores_gemma":[0.0039100163,0.004421912,0.0655428,0.005603912,0.009499705,0.0020548804,0.0026121822,0.7706542,0.023285313,0.019631533,0.09234832,0.0004351889],"about_ca_topic_score_codex":0.011116392,"about_ca_topic_score_gemma":0.021288523,"teacher_disagreement_score":0.9536906,"about_ca_system_score_codex":0.00210114,"about_ca_system_score_gemma":0.0091412505,"threshold_uncertainty_score":0.24491066},"labels":[],"label_agreement":null},{"id":"W4416431165","doi":"10.2196/73822","title":"Identifying Biomedical Entities for Datasets in Scientific Articles: 4-Step Cache-Augmented Generation Approach Using GPT-4o and PubTator 3.0","year":2025,"lang":"en","type":"article","venue":"JMIR Formative Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Metadata; Workflow; Data integration; Metadata modeling; Identification (biology)","score_opus":0.14552605988013695,"score_gpt":0.4481222549381302,"score_spread":0.30259619505799323,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416431165","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.049246576,0.0023978802,0.82123697,0.0022268875,0.0005328559,0.0035313612,0.062265888,0.049535684,0.00902592],"genre_scores_gemma":[0.07739358,0.00048642547,0.8489916,0.0003555154,0.00009968441,0.0024062477,0.066277385,0.0020393971,0.0019502173],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9949366,0.0019095239,0.0007785523,0.0012141578,0.0010119019,0.00014923605],"domain_scores_gemma":[0.96757925,0.020618338,0.0027567546,0.004375264,0.004269532,0.00040088518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011534553,0.0013663144,0.00093339203,0.009876101,0.0011639987,0.0039998204,0.0016825076,0.0013545648,0.009453117],"category_scores_gemma":[0.043994997,0.00081051124,0.003429459,0.005304206,0.00083269016,0.0027246575,0.004804059,0.0013485947,0.0045084762],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013528104,0.00039301225,0.055509925,0.011527909,0.0013886997,0.00244999,0.006169862,0.016588569,0.042639673,0.024832407,0.119914286,0.7172329],"study_design_scores_gemma":[0.0008065783,0.00069932075,0.06373636,0.0024978665,0.003003983,0.004235582,0.0037691484,0.35729772,0.07168408,0.08495892,0.4067307,0.00057972135],"about_ca_topic_score_codex":0.0038884915,"about_ca_topic_score_gemma":0.0114224795,"teacher_disagreement_score":0.011534553,"about_ca_system_score_codex":0.0013941831,"about_ca_system_score_gemma":0.0050520394,"threshold_uncertainty_score":0.0610013},"labels":[],"label_agreement":null},{"id":"W4416437085","doi":"10.48550/arxiv.2511.02824","title":"Kosmos: An AI Scientist for Autonomous Discovery","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"National Institutes of Health; Engineering and Physical Sciences Research Council; UK Dementia Research Institute; Medical Research Council; Cure Alzheimer's Fund; Canadian Institute for Advanced Research; Toyota Research Institute","keywords":"Limiting; Code (set theory); Reading (process); Scientific discovery; Source code","score_opus":0.03738130673923349,"score_gpt":0.33318369300503803,"score_spread":0.2958023862658045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416437085","genre_codex":"methods","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01694567,0.0010861297,0.8563376,0.0057187355,0.0009936459,0.00070684403,0.0023526049,0.08105561,0.0348031],"genre_scores_gemma":[0.11022069,0.001235639,0.8568304,0.0019985237,0.00030678388,0.0010726058,0.004933935,0.0067855874,0.016615808],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99609154,0.0011699483,0.00031406555,0.0008092302,0.0013714524,0.00024374363],"domain_scores_gemma":[0.98558426,0.00619288,0.0007503111,0.003922323,0.002248126,0.0013021962],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0070375362,0.0011361089,0.00081781705,0.0020653768,0.0012166417,0.0043225554,0.0036922917,0.0018233263,0.012769667],"category_scores_gemma":[0.028030492,0.0011955332,0.0015191481,0.0014439973,0.0024531656,0.007248043,0.0068507493,0.003494213,0.008162281],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017177964,0.0005702223,0.009871294,0.0018286335,0.0006446128,0.0013004387,0.0039731306,0.043430515,0.03629002,0.21549901,0.247892,0.43698233],"study_design_scores_gemma":[0.0004952537,0.00032043902,0.0013068335,0.00029154113,0.00017897839,0.00068747567,0.0005685031,0.27351907,0.02327898,0.19766462,0.5015136,0.00017479835],"about_ca_topic_score_codex":0.0013253998,"about_ca_topic_score_gemma":0.0022663516,"teacher_disagreement_score":0.012769667,"about_ca_system_score_codex":0.0010977832,"about_ca_system_score_gemma":0.0039203335,"threshold_uncertainty_score":0.042718828},"labels":[],"label_agreement":null},{"id":"W4416660295","doi":"10.5256/f1000research.26233.r99140","title":"Referee report. For: Methods developed during the first National Center for Biotechnology Information Structural Variation Codeathon at Baylor College of Medicine [version 1; peer review: 1 approved, 1 approved with reservations]","year":2021,"lang":"en","type":"article","venue":"Faculty of 1000 Research Ltd","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Institute of Genetics; U.S. National Library of Medicine; National Institute of Neurological Disorders and Stroke; National Institute of General Medical Sciences; U.S. Department of Health and Human Services; National Institutes of Health; National Cancer Institute; Oxford Nanopore Technologies; Rice University","keywords":"Center (category theory); Information center; Information system; Information technology; Variation (astronomy)","score_opus":0.06898198587702943,"score_gpt":0.4024237387519111,"score_spread":0.3334417528748817,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416660295","genre_codex":"editorial","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00068934093,0.0023050855,0.030610185,0.24903919,0.37680182,0.0038426814,0.19620658,0.01680419,0.12370088],"genre_scores_gemma":[0.0073258197,0.0028135425,0.03701475,0.11420751,0.05791254,0.0061469297,0.09703734,0.015033086,0.6625084],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9890469,0.002505851,0.0013799102,0.001168348,0.005341204,0.00055782654],"domain_scores_gemma":[0.77565116,0.045333873,0.0027649712,0.0121505875,0.15972327,0.0043761553],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.015920587,0.0014655672,0.0020985394,0.0067051114,0.0043043797,0.0042727925,0.003858249,0.0060251257,0.6828353],"category_scores_gemma":[0.2041088,0.0010543334,0.0024458705,0.005658808,0.0015119272,0.0044232113,0.0054160575,0.005654876,0.4901171],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000072787834,0.000008138102,0.000023916726,0.000049144466,0.0000024724627,0.000013881651,0.000019199952,0.000008417374,0.00006094201,0.00018539678,0.9970534,0.0025677579],"study_design_scores_gemma":[0.000055709002,0.000011648411,0.00073410745,0.0002247815,0.000016483671,0.00008184212,0.00015019643,0.00014283785,0.0002899815,0.0018788831,0.9963574,0.00005603889],"about_ca_topic_score_codex":0.036134843,"about_ca_topic_score_gemma":0.054219425,"teacher_disagreement_score":0.9840794,"about_ca_system_score_codex":0.005135665,"about_ca_system_score_gemma":0.007981495,"threshold_uncertainty_score":0.45239693},"labels":[],"label_agreement":null},{"id":"W4416716702","doi":"10.1101/2025.11.23.690073","title":"A Multi-Agent Approach to Generating Context-Rich Gene Sets","year":2025,"lang":"","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Set (abstract data type); Relevance (law); Gene nomenclature; Pipeline (software); Gene; Selection (genetic algorithm); Gene prediction; Task (project management); Biological data","score_opus":0.026603529936153152,"score_gpt":0.2584987157955012,"score_spread":0.23189518585934807,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416716702","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.039160743,0.00022941777,0.9535244,0.0006041034,0.000060921695,0.0002120297,0.0008464034,0.003889364,0.001472578],"genre_scores_gemma":[0.28337553,0.00015476371,0.71248335,0.0002833307,0.000027653477,0.00042729,0.0011522234,0.000245617,0.0018501922],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999433,0.00018117834,0.000043221844,0.00017414201,0.0001286926,0.00003974612],"domain_scores_gemma":[0.99851686,0.0010649655,0.00009598035,0.00011841293,0.00013374658,0.00007001697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011599753,0.00084598875,0.00078775,0.0011297292,0.000811975,0.0011649594,0.0016116044,0.0009972474,0.0025119989],"category_scores_gemma":[0.0037678971,0.0005348527,0.0013654135,0.0007422224,0.00073043065,0.0011926027,0.0017566428,0.0011011012,0.000376293],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019058032,0.00015155604,0.0022353216,0.00024399758,0.00013196067,0.0005820043,0.00037986823,0.8855443,0.008770375,0.017404139,0.0034206929,0.08094529],"study_design_scores_gemma":[0.000024344514,0.000015943657,0.00008953955,0.0000071724426,0.000016861986,0.000026549344,0.00002960949,0.98873645,0.0015428842,0.007838339,0.0016646137,0.000007833801],"about_ca_topic_score_codex":0.006495817,"about_ca_topic_score_gemma":0.01043361,"teacher_disagreement_score":0.006495817,"about_ca_system_score_codex":0.00115111,"about_ca_system_score_gemma":0.0015486983,"threshold_uncertainty_score":0.012915969},"labels":[],"label_agreement":null},{"id":"W4416865968","doi":"10.1093/database/baag027","title":"Large-scale Manual Curation and Harmonization of Metadata from Metagenomic and Cancer Genomic Repositories: Challenges and Solutions","year":2025,"lang":"en","type":"preprint","venue":"Database","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institute of Population and Public Health","funders":"National Cancer Institute; National Institutes of Health","keywords":"Metadata; Data curation; Harmonization; Geospatial metadata; Ontology; Standardization; Metadata repository; Digital curation","score_opus":0.041418253084737164,"score_gpt":0.30703858055245375,"score_spread":0.2656203274677166,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4416865968","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13028455,0.012114221,0.76099837,0.046131935,0.0013599143,0.0048344047,0.019273799,0.012672507,0.012330377],"genre_scores_gemma":[0.094063975,0.002365181,0.8677826,0.0063005155,0.00035682597,0.0028265025,0.02059427,0.003656266,0.0020538496],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.8538506,0.07167219,0.029604154,0.015413058,0.026903903,0.002556139],"domain_scores_gemma":[0.5340558,0.17378949,0.048403278,0.17675652,0.0625301,0.0044648442],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1868279,0.0017463745,0.002622961,0.017641174,0.007815648,0.013874441,0.007950565,0.002304157,0.0023452824],"category_scores_gemma":[0.2648692,0.002207272,0.0029731619,0.022422573,0.0050455746,0.012239569,0.021411825,0.004389629,0.0016723315],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044817844,0.00065602985,0.10446231,0.009345648,0.0020798917,0.0015615881,0.03125447,0.0057486617,0.031675592,0.036162365,0.08165633,0.694949],"study_design_scores_gemma":[0.00029108085,0.0004137144,0.09541553,0.010297843,0.0023025984,0.0032017215,0.027486028,0.022908768,0.065528005,0.15689555,0.61417687,0.001082387],"about_ca_topic_score_codex":0.012063584,"about_ca_topic_score_gemma":0.024506206,"teacher_disagreement_score":0.8131721,"about_ca_system_score_codex":0.0045644194,"about_ca_system_score_gemma":0.0315631,"threshold_uncertainty_score":0.9880522},"labels":[],"label_agreement":null},{"id":"W4417079696","doi":"10.1016/j.phacli.2025.09.372","title":"Utilisation de l’intelligence artificielle générative dans le développement du grand dictionnaire canadien de l’histoire de la pharmacie","year":2025,"lang":"fr","type":"article","venue":"Le Pharmacien Clinicien","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université de Montréal; Centre Hospitalier Universitaire Sainte-Justine","funders":"","keywords":"Web site; Index (typography); Database query","score_opus":0.03433401180321801,"score_gpt":0.34800638549371526,"score_spread":0.3136723736904973,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417079696","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011930855,0.0006522962,0.971144,0.0015623771,0.00017547516,0.00024925114,0.000753041,0.0029257538,0.010606928],"genre_scores_gemma":[0.07934801,0.0006802692,0.91225064,0.0005203399,0.00007944007,0.00016882655,0.0020935887,0.0005421877,0.0043166787],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.991402,0.0034997696,0.001130241,0.0013195918,0.0024683215,0.00018014736],"domain_scores_gemma":[0.97719806,0.016156029,0.0006583059,0.0032097143,0.0025546246,0.00022331321],"candidate_categories":["sts"],"consensus_categories":[],"category_scores_codex":[0.0075441697,0.0010994928,0.0011321887,0.006431005,0.0012942418,0.0072429706,0.0025309103,0.0016455483,0.006175026],"category_scores_gemma":[0.030422678,0.000905916,0.0029842758,0.0033907741,0.0027887153,0.004760926,0.0041462616,0.0030650643,0.002775931],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026102376,0.00027199,0.005515627,0.0018079681,0.00041374797,0.0012490498,0.007423196,0.02538363,0.017682524,0.20708519,0.011313432,0.72159255],"study_design_scores_gemma":[0.00013863138,0.0001929881,0.0058791656,0.0013996867,0.000535373,0.0024802128,0.0025821165,0.31270197,0.041068386,0.23505832,0.397758,0.00020515079],"about_ca_topic_score_codex":0.009030075,"about_ca_topic_score_gemma":0.0098881805,"teacher_disagreement_score":0.99870574,"about_ca_system_score_codex":0.0018345201,"about_ca_system_score_gemma":0.002712909,"threshold_uncertainty_score":0.03989786},"labels":[],"label_agreement":null},{"id":"W4417121609","doi":"10.64898/2025.12.01.691701","title":"Cognitive cartography of mammalian brains using meta-analysis of AI experts","year":2025,"lang":"","type":"article","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Mila - Quebec Artificial Intelligence Institute; McGill University; Montreal Neurological Institute and Hospital","funders":"","keywords":"Cognition; GRASP; Cognitive map; Computational model; Cognitive model; Cognitive neuroscience; Encoding (memory); Perspective (graphical)","score_opus":0.034935588603546466,"score_gpt":0.29267383189440227,"score_spread":0.2577382432908558,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417121609","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13330181,0.02003792,0.8127786,0.005296179,0.00025342093,0.00037276957,0.016611092,0.003996024,0.007352235],"genre_scores_gemma":[0.6179728,0.0044514723,0.36268017,0.00067600876,0.0001347888,0.0006273509,0.011939925,0.000781967,0.00073547533],"study_design_codex":"design_other","study_design_gemma":"meta_analysis","domain_scores_codex":[0.9932637,0.0043728594,0.00036635363,0.0014578144,0.00042197618,0.00011726273],"domain_scores_gemma":[0.9689075,0.02374302,0.0018643887,0.0041117696,0.0010175887,0.00035567262],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020082511,0.0016944621,0.0017226771,0.017865028,0.00094990473,0.006535109,0.0018942902,0.0009988506,0.0024273123],"category_scores_gemma":[0.041443933,0.0008681849,0.008356723,0.009367884,0.0013300754,0.003062383,0.0025550846,0.0016706748,0.00045810323],"study_design_candidate":"meta_analysis","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009932774,0.0002207017,0.22681418,0.009764528,0.086246505,0.0010170497,0.007737787,0.20949799,0.011787267,0.13819377,0.017281543,0.2904454],"study_design_scores_gemma":[0.00010901333,0.00015752432,0.044632476,0.001128798,0.014014436,0.00042542315,0.0019164557,0.31472516,0.004901908,0.5835214,0.03419395,0.00027340907],"about_ca_topic_score_codex":0.005880575,"about_ca_topic_score_gemma":0.009197887,"teacher_disagreement_score":0.020082511,"about_ca_system_score_codex":0.0022817955,"about_ca_system_score_gemma":0.0023500891,"threshold_uncertainty_score":0.10620779},"labels":[],"label_agreement":null},{"id":"W4417194498","doi":"10.22541/au.176536821.14681824/v1","title":"Drug Terminology Server in Nepal: A call for implementation","year":2025,"lang":"","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Wycliffe College","funders":"","keywords":"Terminology; The Internet; Key (lock); Troubleshooting; Action (physics)","score_opus":0.015562588457577443,"score_gpt":0.348688906708369,"score_spread":0.33312631825079153,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417194498","genre_codex":"methods","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019214949,0.00074374833,0.45362046,0.022039589,0.0019864568,0.0023480789,0.030499006,0.43745297,0.032094754],"genre_scores_gemma":[0.11779826,0.0010534977,0.65564275,0.015511766,0.0006142981,0.0020805988,0.10673247,0.056437925,0.044128455],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99576914,0.00065665116,0.0006441279,0.00089337525,0.0013709991,0.0006657146],"domain_scores_gemma":[0.9757297,0.0065102875,0.0008575924,0.008811435,0.005500122,0.0025907878],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013512342,0.0015135991,0.0021136014,0.0023635568,0.0023450325,0.009068199,0.0068792123,0.0044338936,0.037094813],"category_scores_gemma":[0.02693776,0.0025218164,0.0025676116,0.0031743676,0.0017167095,0.013787368,0.00969152,0.0067568356,0.038792837],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002855426,0.0014642845,0.013066737,0.0013983231,0.00056641566,0.0012559388,0.0027058206,0.0017680356,0.028659534,0.03859635,0.5898423,0.31782082],"study_design_scores_gemma":[0.0011777637,0.00040740948,0.010673548,0.000823517,0.0005522155,0.00094460504,0.0015430154,0.052155837,0.046923146,0.040736016,0.8435272,0.00053570967],"about_ca_topic_score_codex":0.013030373,"about_ca_topic_score_gemma":0.010691825,"teacher_disagreement_score":0.037094813,"about_ca_system_score_codex":0.0028723914,"about_ca_system_score_gemma":0.007915124,"threshold_uncertainty_score":0.124094486},"labels":[],"label_agreement":null},{"id":"W4417286561","doi":"10.64898/2025.12.12.693752","title":"React-to-Me: A Conversational Interface for Interactive Exploration of the Reactome Pathway Knowledgebase","year":2025,"lang":"en","type":"article","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Vector Institute; York University; University of Toronto; Ontario Institute for Cancer Research","funders":"National Institutes of Health; Canada First Research Excellence Fund","keywords":"Natural language user interface; Interface (matter); Usability; User interface; Natural language; Reliability (semiconductor); Natural language generation; Natural language understanding","score_opus":0.015021058629709243,"score_gpt":0.2610686276250647,"score_spread":0.2460475689953555,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417286561","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03025307,0.0009865017,0.6396996,0.0010940325,0.00031858237,0.0013625256,0.018574558,0.29349232,0.014218767],"genre_scores_gemma":[0.2927548,0.0014388878,0.60949975,0.002586688,0.00024004445,0.0069704037,0.03328221,0.02794601,0.025281249],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987425,0.00070038345,0.00010472281,0.00024118206,0.0001471733,0.00006406365],"domain_scores_gemma":[0.9938543,0.0051998477,0.00014706349,0.00035719952,0.00021143402,0.00023019468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003488009,0.0021559105,0.0006793779,0.001104344,0.00042819296,0.0012957951,0.0019324577,0.0014575475,0.045170963],"category_scores_gemma":[0.011479961,0.0006519697,0.0010919622,0.00033952767,0.00041658065,0.0027791746,0.004193752,0.0012513013,0.010347427],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.008727953,0.0008915273,0.005130314,0.008865394,0.0005326397,0.004314021,0.023770701,0.0115809785,0.1398971,0.026874619,0.32859638,0.44081837],"study_design_scores_gemma":[0.0013607548,0.00094837416,0.004795228,0.0011291949,0.0003086149,0.00228328,0.003571874,0.18808302,0.07473504,0.076855004,0.64531654,0.00061307865],"about_ca_topic_score_codex":0.000687118,"about_ca_topic_score_gemma":0.0010434953,"teacher_disagreement_score":0.045170963,"about_ca_system_score_codex":0.00036180092,"about_ca_system_score_gemma":0.0006246599,"threshold_uncertainty_score":0.1511119},"labels":[],"label_agreement":null},{"id":"W4417325098","doi":"10.1177/15705838251394800","title":"How Information Complies With a Template: A Dual Mereological System","year":2025,"lang":"en","type":"article","venue":"Applied Ontology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Mereology; Relation (database); Axiom; Extensional definition; Dual (grammatical number); Isomorphism (crystallography); Polyhedron; Ontology","score_opus":0.006717573772361692,"score_gpt":0.21567984689596403,"score_spread":0.20896227312360235,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4417325098","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.040454935,0.000545761,0.8660237,0.005793114,0.00035757423,0.00019886227,0.00042672237,0.0011360283,0.08506329],"genre_scores_gemma":[0.57406914,0.0004040631,0.40751258,0.001561662,0.00043583207,0.00032026143,0.0009075359,0.00043460826,0.014354232],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99134517,0.003106213,0.0009809282,0.0021050626,0.0019412,0.0005214748],"domain_scores_gemma":[0.9887827,0.004629666,0.0008225177,0.0038389731,0.0014623302,0.00046386968],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008953457,0.0005975603,0.00082027283,0.0033964396,0.0035730868,0.008527095,0.0028292937,0.0039027338,0.0061282823],"category_scores_gemma":[0.01335775,0.0012037095,0.0021827465,0.0022653283,0.013103816,0.025037268,0.0062042433,0.0031344143,0.0020648707],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011359704,0.000009833462,0.00018496715,0.00001820824,0.0000061236497,0.000120294666,0.0006234604,0.00026630648,0.0003489287,0.9946404,0.0004117247,0.0033582705],"study_design_scores_gemma":[0.000019698393,0.000025412966,0.00023706576,0.000051740008,0.000031156607,0.0006245577,0.00035347647,0.00718879,0.0011818287,0.9573521,0.03288916,0.00004501599],"about_ca_topic_score_codex":0.002214308,"about_ca_topic_score_gemma":0.0015629423,"teacher_disagreement_score":0.008953457,"about_ca_system_score_codex":0.0022742937,"about_ca_system_score_gemma":0.0022936095,"threshold_uncertainty_score":0.047350943},"labels":[],"label_agreement":null},{"id":"W4531093","doi":"","title":"Extraction of Disease-Treatment Semantic Relations from Biomedical Sentences","year":2010,"lang":"en","type":"article","venue":"Meeting of the Association for Computational Linguistics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Relationship extraction; Relation (database); Semantic relation; Measure (data warehouse); Computer science; Natural language processing; Focus (optics); Artificial intelligence; Semantics (computer science); Information retrieval; Information extraction; Data mining; Medicine; Programming language","score_opus":0.012831567719147929,"score_gpt":0.2901257257106199,"score_spread":0.27729415799147195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4531093","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23160559,0.012894242,0.60674894,0.007939966,0.0015537458,0.0030587341,0.10171468,0.008703025,0.025781082],"genre_scores_gemma":[0.3063249,0.0025894954,0.5997563,0.00081711495,0.0006653175,0.0008578229,0.086673215,0.00029774423,0.0020179912],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9970175,0.0010481251,0.0005859991,0.0006483123,0.00059765665,0.000102406164],"domain_scores_gemma":[0.9887217,0.007994466,0.0012139634,0.0006097175,0.001276883,0.00018317725],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002777044,0.0018594894,0.0009872274,0.0076245246,0.0013546454,0.0015557129,0.0007805753,0.0014487085,0.0045221425],"category_scores_gemma":[0.012331245,0.0005195304,0.0017803541,0.0037628154,0.00055676274,0.0028066065,0.0013160845,0.0015192045,0.002280975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0014800844,0.0010691182,0.033326,0.0104933735,0.0007212496,0.005336471,0.0041742725,0.0097421175,0.1425313,0.033111703,0.045336932,0.7126773],"study_design_scores_gemma":[0.000577191,0.0012612836,0.11265064,0.0025095239,0.0031379326,0.013308503,0.0066016647,0.21402572,0.17376798,0.1312646,0.3404145,0.0004804946],"about_ca_topic_score_codex":0.0019997214,"about_ca_topic_score_gemma":0.0022639614,"teacher_disagreement_score":0.0076245246,"about_ca_system_score_codex":0.0009379207,"about_ca_system_score_gemma":0.002812756,"threshold_uncertainty_score":0.015128076},"labels":[],"label_agreement":null},{"id":"W47286948","doi":"10.1007/978-3-642-22218-4_28","title":"COPE: Childhood Obesity Prevention [Knowledge] Enterprise","year":2011,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":15,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Ontology; Childhood obesity; Obesity; Computer science; Overweight; Domain (mathematical analysis); Work (physics); Domain knowledge; Knowledge management; Discipline; Medicine; Engineering; Social science; Sociology","score_opus":0.015521637088377034,"score_gpt":0.25944627789415314,"score_spread":0.2439246408057761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W47286948","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.008408888,0.009414924,0.18305364,0.009813894,0.0010899005,0.00040805255,0.33948734,0.13548261,0.31284076],"genre_scores_gemma":[0.057817902,0.016654715,0.225409,0.0054802084,0.0008430408,0.0006660564,0.57406276,0.013739099,0.10532731],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99975044,0.000039190007,0.000024636623,0.00005846939,0.00010299286,0.000024236495],"domain_scores_gemma":[0.9991806,0.0003767961,0.000058327827,0.0001691568,0.00010413491,0.00011099172],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051861827,0.00095644745,0.00070011104,0.0023607542,0.00042341163,0.002695115,0.0016154835,0.0013964436,0.053131457],"category_scores_gemma":[0.0026278158,0.00043138143,0.0008473474,0.0036860732,0.00022614289,0.0030546826,0.00231182,0.0011226337,0.029687831],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013815088,0.00006665673,0.0013165048,0.0009807539,0.000081841376,0.0004113089,0.00019997009,0.003341594,0.0014869724,0.015918532,0.6654166,0.3106411],"study_design_scores_gemma":[0.000049750965,0.000016586571,0.0019333005,0.0003312127,0.000070326845,0.00048666686,0.0001097484,0.007367488,0.0018660456,0.027636733,0.96008605,0.00004607039],"about_ca_topic_score_codex":0.003853604,"about_ca_topic_score_gemma":0.0053442037,"teacher_disagreement_score":0.053131457,"about_ca_system_score_codex":0.0005402966,"about_ca_system_score_gemma":0.0011262128,"threshold_uncertainty_score":0.17774242},"labels":[],"label_agreement":null},{"id":"W50348749","doi":"","title":"Recognition of Multi-sentence n-ary Subcellular Localization Mentions in Biomedical Abstracts.","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Subcellular localization; Sentence; Artificial intelligence; Natural language processing; Computer science; Biology; Biochemistry; Gene","score_opus":0.03283894279921447,"score_gpt":0.29869265346320634,"score_spread":0.2658537106639919,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W50348749","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.4000462,0.0046607587,0.54577744,0.002740904,0.0005702764,0.00061655,0.017284095,0.016411034,0.01189273],"genre_scores_gemma":[0.5822532,0.00088853686,0.39462483,0.00023061194,0.000316813,0.0001955236,0.017770456,0.00052009313,0.0031998493],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99881965,0.00023787668,0.00019372095,0.0003983244,0.00029037974,0.000060043094],"domain_scores_gemma":[0.9909889,0.0054571843,0.001673025,0.00046058124,0.0012331425,0.00018726112],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016633556,0.00067959254,0.00043873428,0.004078682,0.00081811124,0.0013040054,0.00080551085,0.0010772384,0.0029582658],"category_scores_gemma":[0.0085847145,0.0002969684,0.00071844354,0.0024736843,0.00049625395,0.0027917763,0.0008026892,0.0006558088,0.0017446663],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013495939,0.00033354523,0.025867261,0.0039464263,0.00028458796,0.005676911,0.005639688,0.008355819,0.22862907,0.019287007,0.033524193,0.6671059],"study_design_scores_gemma":[0.00013838832,0.0007664547,0.08402759,0.0007460426,0.001076478,0.010958096,0.0049839327,0.39918894,0.24333104,0.07406341,0.18040802,0.00031164868],"about_ca_topic_score_codex":0.0019980834,"about_ca_topic_score_gemma":0.0035885742,"teacher_disagreement_score":0.004078682,"about_ca_system_score_codex":0.00068183296,"about_ca_system_score_gemma":0.0008157476,"threshold_uncertainty_score":0.009896398},"labels":[],"label_agreement":null},{"id":"W54772941","doi":"10.1007/978-3-642-38457-8_6","title":"Identifying Explicit Discourse Connectives in Text","year":2013,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":17,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Treebank; Computer science; Parsing; Natural language processing; Artificial intelligence; Syntax; Coherence (philosophical gambling strategy); Head (geology); Linguistics; Mathematics","score_opus":0.023970655486770534,"score_gpt":0.2931898755067489,"score_spread":0.2692192200199784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W54772941","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.2745983,0.005198806,0.6462583,0.005611491,0.0010898765,0.00070636126,0.011817166,0.00533897,0.04938071],"genre_scores_gemma":[0.6840438,0.0019426398,0.2862128,0.00038706142,0.0005400263,0.00044711682,0.0145614585,0.0007762213,0.011088861],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99580765,0.0017863194,0.0004686678,0.0009055074,0.0008322405,0.00019956389],"domain_scores_gemma":[0.9740968,0.021776583,0.0010249089,0.000641735,0.0021550264,0.00030504874],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030768102,0.001267641,0.0006691253,0.0061820936,0.0018228021,0.004425102,0.0012322278,0.0013145474,0.011387554],"category_scores_gemma":[0.019845942,0.00074375956,0.0006213801,0.0041015455,0.0015775679,0.011413856,0.0045522857,0.0018437491,0.0036054286],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001308113,0.00033252462,0.013765291,0.004380372,0.00019174826,0.0041437442,0.031742357,0.0041794055,0.062282566,0.25346017,0.034726676,0.58948714],"study_design_scores_gemma":[0.00021405265,0.00030132764,0.009087578,0.0030343635,0.00075011497,0.0030595555,0.029057888,0.16690248,0.10239558,0.43910897,0.24585956,0.00022850308],"about_ca_topic_score_codex":0.0011358833,"about_ca_topic_score_gemma":0.0014694765,"teacher_disagreement_score":0.011387554,"about_ca_system_score_codex":0.0010791976,"about_ca_system_score_gemma":0.0014524795,"threshold_uncertainty_score":0.038095176},"labels":[],"label_agreement":null},{"id":"W62816693","doi":"","title":"A Design Methodology for a Biomedical Literature Indexing Tool Using the Rhetoric of Science","year":2004,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo; Western University","funders":"","keywords":"Computer science; Search engine indexing; Citation; Context (archaeology); Information retrieval; Science Citation Index; Scientific literature; Function (biology); Data science; Relation (database); Domain (mathematical analysis); Process (computing); Rhetorical question; Index (typography); World Wide Web; Data mining; Mathematics","score_opus":0.11300373425592718,"score_gpt":0.37588052810570693,"score_spread":0.26287679384977974,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W62816693","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0018456207,0.000035358426,0.99408805,0.00037432587,0.0000563652,0.0011143363,0.00007523115,0.0013678443,0.0010428737],"genre_scores_gemma":[0.009963336,0.00003588645,0.9870504,0.0001049819,0.000030797117,0.0014356067,0.00014203564,0.00015279365,0.0010842168],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9874708,0.0053498126,0.001982444,0.0018344265,0.0030505813,0.00031188593],"domain_scores_gemma":[0.96573114,0.017944327,0.0028483346,0.0041020163,0.008194646,0.0011796086],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.021042455,0.0013116547,0.0008813735,0.0065149767,0.0022822134,0.006235504,0.0027797655,0.0025506904,0.006792196],"category_scores_gemma":[0.043069925,0.001268803,0.0015663048,0.0035119706,0.00384052,0.007419033,0.0035402556,0.002243375,0.0034994707],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004441977,0.00089726935,0.0039651236,0.0026978927,0.00021674155,0.0013680777,0.013226325,0.0065459427,0.063785605,0.33263135,0.011445174,0.5627763],"study_design_scores_gemma":[0.00062231347,0.0017497429,0.0025784536,0.0008527143,0.0006089339,0.0040605776,0.0043307007,0.1897224,0.11673073,0.23649539,0.44177678,0.00047120106],"about_ca_topic_score_codex":0.00061417976,"about_ca_topic_score_gemma":0.00081187865,"teacher_disagreement_score":0.021042455,"about_ca_system_score_codex":0.0015796444,"about_ca_system_score_gemma":0.0040654684,"threshold_uncertainty_score":0.111284435},"labels":[],"label_agreement":null},{"id":"W652033195","doi":"","title":"Expert-novice differences in mammogram interpretation","year":2007,"lang":"en","type":"article","venue":"eScholarship (California Digital Library)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada; Concordia University; University of Pittsburgh; Georgia Institute of Technology","keywords":"Memphis; Interpretation (philosophy); Cognition; Psychology; Medical education; Applied psychology; Medicine; Computer science; Psychiatry","score_opus":0.013792452521885296,"score_gpt":0.2467663118862741,"score_spread":0.2329738593643888,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W652033195","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9959525,0.00019110605,0.0015551526,0.00004469522,0.000008485243,0.000022724844,0.00004728908,0.000028424838,0.00214961],"genre_scores_gemma":[0.9982375,0.000089136294,0.000862035,0.00005009524,0.000007703102,0.000014508588,0.00009463891,0.000015027097,0.0006294237],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99638355,0.0012163023,0.00044698047,0.00054410385,0.0011351638,0.00027395668],"domain_scores_gemma":[0.90270585,0.07861857,0.0065124575,0.0040891776,0.005549681,0.0025241487],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004883359,0.00028281222,0.00026847285,0.0015288055,0.00020779512,0.0012095482,0.00042244585,0.0004714747,0.0026912868],"category_scores_gemma":[0.069113836,0.00022400956,0.00028677314,0.00027726172,0.0006039341,0.001151558,0.0011367823,0.0006098578,0.00041207724],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019639777,0.001406235,0.7728086,0.0004635984,0.00054749264,0.0012238701,0.023834806,0.0035567756,0.02700928,0.0008242778,0.0013050822,0.16505598],"study_design_scores_gemma":[0.00006280907,0.0015341786,0.97630435,0.00006314745,0.000070012866,0.0022078198,0.005629181,0.004072486,0.005723613,0.0019239848,0.0023310343,0.000077333876],"about_ca_topic_score_codex":0.000514685,"about_ca_topic_score_gemma":0.0008257488,"teacher_disagreement_score":0.004883359,"about_ca_system_score_codex":0.00027447796,"about_ca_system_score_gemma":0.00019362867,"threshold_uncertainty_score":0.025825977},"labels":[],"label_agreement":null},{"id":"W66296548","doi":"","title":"Automated Extraction of Protein Mutation Impacts from the Biomedical Literature","year":2011,"lang":"en","type":"article","venue":"Spectrum Research Repository (Concordia University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Concordia University","funders":"","keywords":"Computer science; Information extraction; Information retrieval; Mutation; Ontology; Organism; Precision and recall; Focus (optics); Computational biology; Data mining; Gene; Biology; Genetics","score_opus":0.029735259764050132,"score_gpt":0.29056421128958737,"score_spread":0.26082895152553726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W66296548","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.12428692,0.022535862,0.2676882,0.003427612,0.0011198538,0.0018593939,0.35143214,0.20562202,0.022028038],"genre_scores_gemma":[0.13286242,0.009156604,0.47718397,0.00077933015,0.00074667216,0.0007071883,0.36739594,0.0036847778,0.0074830605],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9973679,0.00023225902,0.0005318028,0.0008982587,0.0008597089,0.00011000002],"domain_scores_gemma":[0.9913662,0.0034943724,0.00163406,0.0009997614,0.0021746117,0.0003309202],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0022141826,0.0021865207,0.0012524979,0.034354147,0.0011052022,0.002960994,0.001324642,0.0011635215,0.0065812557],"category_scores_gemma":[0.010074847,0.0007386876,0.0018663885,0.010219159,0.00057788513,0.0032652838,0.0026452723,0.00096465787,0.006608831],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006501713,0.00029540164,0.031012645,0.008849694,0.0006891215,0.004724047,0.0013904808,0.0032754738,0.08193088,0.0044996464,0.13071597,0.73196656],"study_design_scores_gemma":[0.00018043336,0.000541168,0.14648017,0.00212944,0.002058256,0.0080846585,0.0014614847,0.0725116,0.13531905,0.017756943,0.61305946,0.0004172948],"about_ca_topic_score_codex":0.003114357,"about_ca_topic_score_gemma":0.0055584824,"teacher_disagreement_score":0.034354147,"about_ca_system_score_codex":0.0011437339,"about_ca_system_score_gemma":0.0029193594,"threshold_uncertainty_score":0.022016466},"labels":[],"label_agreement":null},{"id":"W66888145","doi":"10.22230/jripe.2009v1n1a7","title":"JRIPE, a new journal with a distinct function, content, and identity","year":2009,"lang":"en","type":"article","venue":"Journal of Research in Interprofessional Practice and Education","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"University of British Columbia; Université de Sherbrooke","funders":"","keywords":"Identity (music); Content (measure theory); Function (biology); Biology; Mathematics; Art; Evolutionary biology; Aesthetics","score_opus":0.07749718247933955,"score_gpt":0.4693272878948675,"score_spread":0.39183010541552793,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W66888145","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019490167,0.10388168,0.03961407,0.33139908,0.3168103,0.0006687875,0.010178711,0.007282066,0.17067516],"genre_scores_gemma":[0.098701864,0.124562,0.11386702,0.089495316,0.21416987,0.000733737,0.01676267,0.008344702,0.33336282],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.98844945,0.001852209,0.0023725464,0.0010107579,0.0058225035,0.0004925853],"domain_scores_gemma":[0.9064351,0.03655658,0.008383089,0.0073694955,0.018037686,0.023218146],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.0125009725,0.00069755164,0.0013087909,0.015223696,0.0021313182,0.024075165,0.0015186497,0.0023592554,0.0208923],"category_scores_gemma":[0.057663936,0.0005163397,0.0006878223,0.011265112,0.0042857756,0.0081479205,0.003135131,0.0061658733,0.009713763],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012166896,0.00011104042,0.0019110345,0.0032414193,0.00011026711,0.00060489686,0.0008708388,0.00012360478,0.003064492,0.025707666,0.75776315,0.20636998],"study_design_scores_gemma":[0.000014553823,0.000031788986,0.0026186078,0.00046444274,0.00004980762,0.0009362914,0.00052142196,0.0001240417,0.0004792175,0.0035647217,0.9911629,0.00003219461],"about_ca_topic_score_codex":0.00063405174,"about_ca_topic_score_gemma":0.002868109,"teacher_disagreement_score":0.97592485,"about_ca_system_score_codex":0.0021562797,"about_ca_system_score_gemma":0.012974843,"threshold_uncertainty_score":0.06989169},"labels":[],"label_agreement":null},{"id":"W6886918351","doi":"10.15468/dl.2kb3pd","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6886918351","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008637056,0.000045790417,0.00007298707,0.000051810748,0.000014379588,0.000008504898,0.9981158,0.0007447337,0.00085967686],"genre_scores_gemma":[0.00018050159,0.000038251543,0.00027112474,0.000044508786,0.0000027189867,0.00003695169,0.9988707,0.00013879426,0.000416497],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99889344,0.00014890876,0.00015307973,0.00038316147,0.00026179425,0.00015952802],"domain_scores_gemma":[0.9978587,0.0005826205,0.0002040144,0.0005572246,0.00052676024,0.00027075756],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009716543,0.0024552809,0.001531989,0.0048110425,0.0010476275,0.0024087317,0.0031172154,0.0020751148,0.09903262],"category_scores_gemma":[0.005593653,0.0008948044,0.0012904567,0.008891044,0.00047913965,0.002216184,0.0024950504,0.0019385102,0.15193881],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041913187,0.000017959886,0.0004452922,0.0005586246,0.000017009108,0.000022461443,0.000025825613,0.00017959725,0.00015635588,0.00047288925,0.99619293,0.0018691686],"study_design_scores_gemma":[0.000091141286,0.000012030208,0.0018352618,0.00017203881,0.00001581516,0.00006196097,0.00007638641,0.00030102406,0.00028348394,0.0009335315,0.99619734,0.000020061028],"about_ca_topic_score_codex":0.024968777,"about_ca_topic_score_gemma":0.042969853,"teacher_disagreement_score":0.90096736,"about_ca_system_score_codex":0.0017934206,"about_ca_system_score_gemma":0.0026499424,"threshold_uncertainty_score":0.3312971},"labels":[],"label_agreement":null},{"id":"W6886965399","doi":"10.15468/dl.gh7zwm","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Identification (biology); Column (typography)","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6886965399","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000073345436,0.000043141496,0.000072242416,0.0000533046,0.000013419235,0.000007678941,0.99832886,0.0006141604,0.0007938058],"genre_scores_gemma":[0.00017396604,0.00003662474,0.00025004256,0.00004396126,0.0000025222027,0.000035559075,0.99894303,0.00012545452,0.00038879606],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998835,0.00014825082,0.0001555283,0.00040151767,0.00028990416,0.00016981713],"domain_scores_gemma":[0.99779594,0.0005968458,0.00020758121,0.00059928297,0.0005200892,0.000280319],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010747802,0.0023590827,0.0016557603,0.0049425913,0.0011608463,0.0025870195,0.0032646754,0.0023735655,0.102422945],"category_scores_gemma":[0.0056901923,0.0009678067,0.0012687667,0.008815862,0.000527447,0.0024285852,0.0027466163,0.0022650775,0.15764621],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040279356,0.000018087618,0.00047084808,0.0006112937,0.00001837745,0.000025178331,0.000030584266,0.0001856464,0.00017869895,0.000533798,0.9960752,0.00181198],"study_design_scores_gemma":[0.00007945207,0.000008733083,0.0018723366,0.00018784909,0.00001547668,0.000059554248,0.00008322813,0.00025013043,0.00028448293,0.00088299357,0.9962568,0.000018955447],"about_ca_topic_score_codex":0.024273517,"about_ca_topic_score_gemma":0.041902747,"teacher_disagreement_score":0.89757705,"about_ca_system_score_codex":0.0019899425,"about_ca_system_score_gemma":0.0027455848,"threshold_uncertainty_score":0.3426389},"labels":[],"label_agreement":null},{"id":"W6886968334","doi":"10.15468/dl.kmd7kc","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6886968334","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008191469,0.00004606893,0.000074711315,0.0000521299,0.000014282217,0.000008568877,0.99813336,0.00073495606,0.0008539733],"genre_scores_gemma":[0.00017991207,0.000040298197,0.00027571985,0.00004662322,0.0000027517553,0.00003779738,0.9988589,0.00013870465,0.0004192576],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99893683,0.00014259385,0.00014686232,0.0003710107,0.00025022135,0.0001524897],"domain_scores_gemma":[0.99793136,0.0005743798,0.00019685233,0.00053540384,0.0004996203,0.00026235578],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00094327313,0.002425566,0.0015120757,0.004726392,0.0010498005,0.0024006525,0.0031320313,0.002102018,0.10188185],"category_scores_gemma":[0.005534451,0.00089198985,0.0012879113,0.008699142,0.00047627042,0.002248629,0.0024890217,0.0019514052,0.15320504],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040987114,0.000017426535,0.00042594035,0.0005725349,0.000017201432,0.000022889715,0.000025803378,0.00018629778,0.00015426433,0.00048391303,0.9961957,0.0018570427],"study_design_scores_gemma":[0.00009027518,0.000011644796,0.0017225501,0.00017577024,0.000016037484,0.00006168656,0.000074253316,0.00030719783,0.00027687478,0.0009599077,0.9962837,0.000020150897],"about_ca_topic_score_codex":0.0247878,"about_ca_topic_score_gemma":0.04343388,"teacher_disagreement_score":0.89811814,"about_ca_system_score_codex":0.001810897,"about_ca_system_score_gemma":0.0026468867,"threshold_uncertainty_score":0.34082872},"labels":[],"label_agreement":null},{"id":"W6887346995","doi":"10.15468/dl.ybxdfg","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6887346995","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008204781,0.00004505214,0.0000738137,0.000051991512,0.000013765256,0.00000831573,0.9981273,0.0007292359,0.0008686214],"genre_scores_gemma":[0.00018140553,0.000039914543,0.00027095515,0.00004571793,0.0000026732357,0.00003696713,0.9988631,0.00014085662,0.00041826934],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989248,0.00014548523,0.00014812396,0.0003718553,0.0002540373,0.00015571267],"domain_scores_gemma":[0.9978892,0.0005841781,0.00020180263,0.00054693554,0.0005130057,0.00026487728],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009580919,0.0024206063,0.00149854,0.004833854,0.0010687882,0.002401956,0.0031180854,0.0020724563,0.100996844],"category_scores_gemma":[0.0055772117,0.00090481737,0.001276139,0.009014693,0.00048374603,0.0022532928,0.0024712144,0.0019514943,0.15154487],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041268235,0.00001703518,0.00043495154,0.00056050735,0.000017032728,0.000022735472,0.000026533557,0.00018311366,0.0001544612,0.0004963824,0.99620545,0.0018404323],"study_design_scores_gemma":[0.00008691823,0.000011060451,0.0017412791,0.00017197982,0.00001560836,0.000059834503,0.00007536702,0.00029040244,0.00027299134,0.0009488376,0.99630606,0.000019594958],"about_ca_topic_score_codex":0.026077224,"about_ca_topic_score_gemma":0.04490657,"teacher_disagreement_score":0.89900315,"about_ca_system_score_codex":0.0018434199,"about_ca_system_score_gemma":0.002709013,"threshold_uncertainty_score":0.33786815},"labels":[],"label_agreement":null},{"id":"W6888466618","doi":"10.20381/ruor-4387","title":"Annotation Concept Synthesis and Enrichment Analysis: a Logic-Based Approach to the Interpretation of High-Throughput Biological Experiments","year":2011,"lang":"en","type":"article","venue":"Library and Archives Canada (Government of Canada)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Annotation; Set (abstract data type); Gene Annotation; Interpretation (philosophy); Object (grammar); Process (computing); Temporal annotation; Expression (computer science)","score_opus":0.01118718469461945,"score_gpt":0.18480848391771187,"score_spread":0.1736212992230924,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6888466618","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010854625,0.00018666577,0.9969476,0.00013593855,0.000025682375,0.00012466207,0.0001232057,0.0007049901,0.0006658852],"genre_scores_gemma":[0.022648059,0.0002785878,0.9752764,0.00022536452,0.000056368164,0.00036249627,0.0003864616,0.00012620662,0.0006400392],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99187285,0.0029055004,0.0006286099,0.0014638357,0.002787949,0.00034125618],"domain_scores_gemma":[0.9873372,0.008872578,0.0008052583,0.0010141096,0.0017180922,0.0002527368],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.010805552,0.0025641348,0.0017269611,0.0065082577,0.0013865742,0.0045160116,0.0041438187,0.0013230061,0.0032061646],"category_scores_gemma":[0.012080065,0.0008358716,0.0039821453,0.0038184694,0.0032088323,0.0041457415,0.0030682015,0.0032168005,0.0012133792],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006633825,0.0005309369,0.0023476623,0.0030551995,0.0005835935,0.0014052115,0.001521933,0.06851115,0.06743824,0.383889,0.005675502,0.46437818],"study_design_scores_gemma":[0.00009089251,0.00025310577,0.0008276343,0.00038724867,0.00032944477,0.000630281,0.0002959916,0.46926573,0.05733719,0.438695,0.031713437,0.00017406706],"about_ca_topic_score_codex":0.0018912571,"about_ca_topic_score_gemma":0.0015165551,"teacher_disagreement_score":0.98919445,"about_ca_system_score_codex":0.002474773,"about_ca_system_score_gemma":0.00374725,"threshold_uncertainty_score":0.057145894},"labels":[],"label_agreement":null},{"id":"W6891708803","doi":"10.48448/ejg5-gc86","title":"Generating Extractive and Abstractive Summaries in Parallel from Scientific Articles Incorporating Citing Statements","year":2023,"lang":"en","type":"other","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University","funders":"","keywords":"Feature (linguistics); Computational linguistics; Semantics (computer science); Natural language","score_opus":0.06463724054626332,"score_gpt":0.36238112353217705,"score_spread":0.2977438829859137,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6891708803","genre_codex":"methods","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08821486,0.0030580752,0.75782645,0.0026219003,0.0012253901,0.0012938712,0.059264828,0.06915704,0.017337557],"genre_scores_gemma":[0.13792956,0.0013448097,0.770049,0.00019365635,0.0006812691,0.00045399534,0.07617603,0.0026383332,0.010533393],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981255,0.00041221327,0.0002948094,0.00042311053,0.0006679601,0.000076397366],"domain_scores_gemma":[0.9887121,0.006393818,0.0007883668,0.0009881137,0.002795687,0.00032190053],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023894385,0.002516822,0.0011078513,0.010524417,0.0009912102,0.0031569009,0.0011220975,0.0012622223,0.010657825],"category_scores_gemma":[0.01777025,0.00069769955,0.0016806321,0.006300867,0.00044709182,0.0033548127,0.0026363176,0.0011229749,0.008195305],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009822622,0.00042259708,0.005203811,0.0034631481,0.00067073124,0.0016210156,0.0015091525,0.0125833135,0.04110319,0.009463324,0.09521615,0.82776135],"study_design_scores_gemma":[0.00056446414,0.00078383624,0.012697754,0.0006951614,0.0020740912,0.0019700981,0.0027390998,0.46207753,0.13981967,0.10370703,0.27252987,0.00034133688],"about_ca_topic_score_codex":0.002020871,"about_ca_topic_score_gemma":0.0051622065,"teacher_disagreement_score":0.010657825,"about_ca_system_score_codex":0.0005700209,"about_ca_system_score_gemma":0.0021289173,"threshold_uncertainty_score":0.03565395},"labels":[],"label_agreement":null},{"id":"W6893407171","doi":"10.5281/zenodo.16815812","title":"DEPRECATED: Matcher (Version 1) for Automated Task Alignment in the Genomic API for Model Evaluation (GAME)","year":2025,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Task (project management); Container (type theory); Python (programming language); Transcription (linguistics); Matching (statistics); Ontology alignment; Key (lock); Graph","score_opus":0.03467398594823993,"score_gpt":0.2936418968347249,"score_spread":0.25896791088648496,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6893407171","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020997785,0.0001366868,0.14685097,0.00018906126,0.00018457866,0.00022668038,0.028587347,0.81294715,0.008777835],"genre_scores_gemma":[0.064036995,0.00033221007,0.26763952,0.0012769946,0.00010887931,0.002134121,0.20855896,0.41565096,0.040261455],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99847835,0.00020524167,0.00012764263,0.0004692906,0.0004981491,0.00022135902],"domain_scores_gemma":[0.99886286,0.0002906026,0.00007679419,0.0004112375,0.0002623866,0.00009608313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021514385,0.0028720528,0.0011559728,0.0013338869,0.0008228253,0.0027553588,0.0041470705,0.0016134628,0.14436023],"category_scores_gemma":[0.0075878655,0.0017854628,0.0021369846,0.0010457783,0.00051435566,0.0041095996,0.0049116686,0.00284123,0.09741427],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00082656206,0.00014947349,0.0024687895,0.0007869608,0.0001791374,0.00016475674,0.00037507206,0.004432592,0.0059318175,0.013287394,0.90195274,0.06944471],"study_design_scores_gemma":[0.00055265945,0.00020092637,0.0053281384,0.00029356702,0.00009209226,0.00038000525,0.00028804547,0.15174922,0.03841602,0.042581365,0.7597972,0.00032067884],"about_ca_topic_score_codex":0.008201265,"about_ca_topic_score_gemma":0.007901081,"teacher_disagreement_score":0.14436023,"about_ca_system_score_codex":0.0017842216,"about_ca_system_score_gemma":0.0018734596,"threshold_uncertainty_score":0.4829331},"labels":[],"label_agreement":null},{"id":"W6893682764","doi":"10.5281/zenodo.3665746","title":"CanDIG CHORD: Canadian Health Omics Repository, Distributed","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Software; Identification (biology); Software development; MEDLINE; Health data","score_opus":0.030924283322395836,"score_gpt":0.24732557746338868,"score_spread":0.21640129414099285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6893682764","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001002863,0.0017646982,0.021683812,0.0035340409,0.0007327607,0.00045533522,0.8443059,0.09247905,0.034041476],"genre_scores_gemma":[0.006043878,0.0021084757,0.042750634,0.0015026515,0.00019075304,0.0006385269,0.9051465,0.017185744,0.02443291],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974618,0.00021282428,0.0002624157,0.00040896528,0.001367004,0.00028701688],"domain_scores_gemma":[0.9835322,0.0023477552,0.0005034444,0.0035640027,0.007909602,0.002143088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057980474,0.0026487568,0.0020290937,0.014168856,0.003738479,0.008021584,0.0073976372,0.0020295263,0.11688346],"category_scores_gemma":[0.024605935,0.0014801775,0.0016616272,0.021253044,0.0011314102,0.005099223,0.006203244,0.002672569,0.08228856],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019779059,0.000014474734,0.00059182325,0.00040290522,0.00005372131,0.000083205035,0.00009673747,0.00023062111,0.0005283104,0.0031092027,0.9657958,0.028895462],"study_design_scores_gemma":[0.00018346866,0.000011360619,0.002171806,0.00027522107,0.00007404537,0.000107040105,0.00011594862,0.0011094615,0.00095127383,0.0038062187,0.9911014,0.0000928312],"about_ca_topic_score_codex":0.63035935,"about_ca_topic_score_gemma":0.68255556,"teacher_disagreement_score":0.36964065,"about_ca_system_score_codex":0.011434719,"about_ca_system_score_gemma":0.056791063,"threshold_uncertainty_score":0.74363506},"labels":[],"label_agreement":null},{"id":"W6901584990","doi":"10.60692/n2cav-c9590","title":"Introducing Explorer of Taxon Concepts with a case study on spider measurement matrix building","year":2016,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"","keywords":"Spider; Set (abstract data type); Matrix (chemical analysis); Normalization (sociology); Pipeline (software); Bridging (networking); Software","score_opus":0.055492109391758485,"score_gpt":0.27589759454271756,"score_spread":0.2204054851509591,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901584990","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5255331,0.0011876088,0.3971275,0.0049329554,0.00047124358,0.0015799027,0.009450262,0.010990332,0.048727114],"genre_scores_gemma":[0.3825651,0.00052779267,0.5927072,0.00039075306,0.00010357294,0.0006495116,0.006272244,0.0024937855,0.014290101],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99747086,0.0012435447,0.00018302463,0.0003624237,0.000647557,0.00009268544],"domain_scores_gemma":[0.9857635,0.010889149,0.00061507995,0.0011229308,0.0011288362,0.00048050124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033169691,0.0005413149,0.00026186055,0.0039026984,0.0014994964,0.002352014,0.0011074133,0.0010466415,0.009246728],"category_scores_gemma":[0.012451359,0.00032431402,0.00073428836,0.0029898398,0.0012392115,0.0033467251,0.0025414592,0.0010208548,0.0018330946],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005368017,0.00073496404,0.0416763,0.0039463644,0.00010501746,0.041293718,0.18019868,0.0090714805,0.06524394,0.061557405,0.07599904,0.51963633],"study_design_scores_gemma":[0.00007782472,0.00026590083,0.024812568,0.0013444579,0.00010131215,0.014728803,0.053039797,0.04948533,0.045554157,0.025761042,0.78459287,0.00023591572],"about_ca_topic_score_codex":0.0031703643,"about_ca_topic_score_gemma":0.00759647,"teacher_disagreement_score":0.009246728,"about_ca_system_score_codex":0.0011175823,"about_ca_system_score_gemma":0.0011370996,"threshold_uncertainty_score":0.03093344},"labels":[],"label_agreement":null},{"id":"W6901752456","doi":"10.60692/tdada-5v222","title":"Introducing Explorer of Taxon Concepts with a case study on spider measurement matrix building","year":2016,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Agriculture and Agri-Food Canada","funders":"","keywords":"Spider; Set (abstract data type); Matrix (chemical analysis); Normalization (sociology); Pipeline (software); Bridging (networking); Software","score_opus":0.055492109391758485,"score_gpt":0.27589759454271756,"score_spread":0.2204054851509591,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901752456","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5255331,0.0011876088,0.3971275,0.0049329554,0.00047124358,0.0015799027,0.009450262,0.010990332,0.048727114],"genre_scores_gemma":[0.3825651,0.00052779267,0.5927072,0.00039075306,0.00010357294,0.0006495116,0.006272244,0.0024937855,0.014290101],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99747086,0.0012435447,0.00018302463,0.0003624237,0.000647557,0.00009268544],"domain_scores_gemma":[0.9857635,0.010889149,0.00061507995,0.0011229308,0.0011288362,0.00048050124],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033169691,0.0005413149,0.00026186055,0.0039026984,0.0014994964,0.002352014,0.0011074133,0.0010466415,0.009246728],"category_scores_gemma":[0.012451359,0.00032431402,0.00073428836,0.0029898398,0.0012392115,0.0033467251,0.0025414592,0.0010208548,0.0018330946],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005368017,0.00073496404,0.0416763,0.0039463644,0.00010501746,0.041293718,0.18019868,0.0090714805,0.06524394,0.061557405,0.07599904,0.51963633],"study_design_scores_gemma":[0.00007782472,0.00026590083,0.024812568,0.0013444579,0.00010131215,0.014728803,0.053039797,0.04948533,0.045554157,0.025761042,0.78459287,0.00023591572],"about_ca_topic_score_codex":0.0031703643,"about_ca_topic_score_gemma":0.00759647,"teacher_disagreement_score":0.009246728,"about_ca_system_score_codex":0.0011175823,"about_ca_system_score_gemma":0.0011370996,"threshold_uncertainty_score":0.03093344},"labels":[],"label_agreement":null},{"id":"W6901818360","doi":"10.60692/bp2mc-br705","title":"Developing data interoperability using standards: A wheat community use case","year":2017,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"","keywords":"Interoperability; Semantic interoperability; Metadata; Ontology; Data exchange; Cross-domain interoperability; Linked data; The Internet","score_opus":0.25977750430107655,"score_gpt":0.3475690416771639,"score_spread":0.08779153737608736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901818360","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32850078,0.002670652,0.4733768,0.07155151,0.00048117008,0.0024225817,0.0006161435,0.0015917822,0.11878862],"genre_scores_gemma":[0.5491925,0.0022921928,0.42561385,0.0056220596,0.00017790812,0.0010505648,0.0018127601,0.0007817424,0.013456447],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9055215,0.060009424,0.0066290437,0.0040088794,0.020120548,0.0037106303],"domain_scores_gemma":[0.885159,0.06466247,0.0038715077,0.018248875,0.024472766,0.0035853875],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.11503063,0.0010130208,0.0009096538,0.0054762354,0.009394334,0.015989417,0.0054030577,0.010764041,0.0025447637],"category_scores_gemma":[0.07486249,0.0013357805,0.0021118643,0.008594692,0.007371901,0.034058988,0.018603228,0.0067604496,0.00080944377],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029363154,0.0019997167,0.035354186,0.0014887478,0.00017252451,0.014348539,0.11205558,0.0075035384,0.009099288,0.51919466,0.028952334,0.26953727],"study_design_scores_gemma":[0.00022810744,0.00084500434,0.008200962,0.0025370943,0.0001882214,0.006300985,0.09732672,0.03898334,0.019799333,0.14041045,0.6848069,0.00037292764],"about_ca_topic_score_codex":0.017378792,"about_ca_topic_score_gemma":0.017321907,"teacher_disagreement_score":0.88496935,"about_ca_system_score_codex":0.008208195,"about_ca_system_score_gemma":0.00940439,"threshold_uncertainty_score":0.6083474},"labels":[],"label_agreement":null},{"id":"W6901844693","doi":"10.6084/m9.figshare.19759383","title":"Additional file 3 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test)","score_opus":0.06651094892406054,"score_gpt":0.4052983664027775,"score_spread":0.338787417478717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901844693","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002848359,0.000014343533,0.0009517929,0.00010855453,0.000023196328,0.00009266124,0.9961677,0.0014929208,0.00086392136],"genre_scores_gemma":[0.0061974907,0.000077588564,0.009313242,0.00031999193,0.00009056983,0.0011563551,0.9763431,0.0013620872,0.0051396275],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9990785,0.00013410773,0.00024003493,0.00024087721,0.00022203497,0.000084447944],"domain_scores_gemma":[0.98294026,0.011836558,0.0012496228,0.0009770559,0.0026081125,0.00038838587],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001610309,0.0010803692,0.0007746469,0.003381597,0.0005663586,0.0015177715,0.0014765777,0.0009171269,0.6729432],"category_scores_gemma":[0.018463261,0.00045347246,0.00090835465,0.0035535928,0.00025361346,0.0013500636,0.001313271,0.0007326297,0.1180505],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038950273,0.00009205604,0.0028794166,0.002437359,0.00006351767,0.00014936471,0.00008014758,0.00049578055,0.00063574873,0.00074951525,0.9740831,0.017944569],"study_design_scores_gemma":[0.002199735,0.00020658971,0.02405914,0.0014742648,0.00021482444,0.00069758086,0.0003627542,0.0045486535,0.0059243334,0.008897787,0.95124906,0.00016525632],"about_ca_topic_score_codex":0.004108138,"about_ca_topic_score_gemma":0.0065211216,"teacher_disagreement_score":0.6729432,"about_ca_system_score_codex":0.0010615097,"about_ca_system_score_gemma":0.0020312287,"threshold_uncertainty_score":0.46650684},"labels":[],"label_agreement":null},{"id":"W6901906532","doi":"10.6084/m9.figshare.12799365","title":"Additional file 3 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2021,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Documentation; Rare events","score_opus":0.05121493344011546,"score_gpt":0.32572751215804435,"score_spread":0.2745125787179289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901906532","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008901042,0.00007236454,0.00070557254,0.00020745854,0.00003019392,0.0005065889,0.9970655,0.00038425523,0.0009390778],"genre_scores_gemma":[0.0045947437,0.00077534886,0.029339358,0.0019480088,0.0002107669,0.023585193,0.9258939,0.0018402211,0.011812465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99626476,0.0009516968,0.0014146002,0.0006627695,0.00048447694,0.00022170082],"domain_scores_gemma":[0.8921012,0.09094481,0.0062795654,0.003198,0.006246082,0.0012303694],"candidate_categories":["metaresearch","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0081131775,0.0012668817,0.001841197,0.008411752,0.00089837535,0.0026006522,0.0020116447,0.0016588287,0.77455777],"category_scores_gemma":[0.08511096,0.0009813752,0.002512411,0.009701538,0.0005092184,0.002585809,0.00239626,0.0014314831,0.09563002],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042649359,0.000046355544,0.0013427264,0.04466585,0.00028185258,0.00007542654,0.00016752457,0.0004386585,0.00024391935,0.002219524,0.93413883,0.015952915],"study_design_scores_gemma":[0.003629256,0.00012571755,0.0070801587,0.016789561,0.00081847137,0.00025288577,0.00030437156,0.0006466138,0.00079251,0.014206558,0.955212,0.00014194199],"about_ca_topic_score_codex":0.004817988,"about_ca_topic_score_gemma":0.013929324,"teacher_disagreement_score":0.99798834,"about_ca_system_score_codex":0.0016435151,"about_ca_system_score_gemma":0.0063235867,"threshold_uncertainty_score":0.32156593},"labels":[],"label_agreement":null},{"id":"W6901967445","doi":"10.6084/m9.figshare.12799362.v2","title":"Additional file 2 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2021,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Disease; Rare events","score_opus":0.0504202911666707,"score_gpt":0.3255831022755051,"score_spread":0.2751628111088344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6901967445","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00009401381,0.00010598268,0.0005278739,0.00020400817,0.000027278396,0.00039255334,0.99741817,0.00036681045,0.0008632394],"genre_scores_gemma":[0.00454858,0.0009984785,0.023377394,0.0018352075,0.0002027333,0.016971212,0.9400198,0.0017095887,0.010336977],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970356,0.0007456552,0.001118262,0.0005345648,0.0003762661,0.00018952263],"domain_scores_gemma":[0.9097334,0.07565419,0.0056690024,0.002558743,0.0052251425,0.0011594697],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006236715,0.0012268733,0.0019873013,0.009423827,0.0008518359,0.0023936399,0.0018741533,0.0015980371,0.75149596],"category_scores_gemma":[0.066778585,0.0008688528,0.0023878135,0.010654296,0.0004713089,0.0025234397,0.0023248561,0.0012994077,0.086941466],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042716903,0.000049928927,0.0013583608,0.06318174,0.00030385822,0.00008912907,0.00016266716,0.0003877754,0.000316053,0.0020857041,0.9151549,0.016482767],"study_design_scores_gemma":[0.0034698474,0.00014093402,0.008048594,0.020075906,0.00095346716,0.00029814168,0.00029596125,0.00062398403,0.00080140436,0.012275011,0.95287377,0.00014299825],"about_ca_topic_score_codex":0.004631547,"about_ca_topic_score_gemma":0.0147834625,"teacher_disagreement_score":0.99376327,"about_ca_system_score_codex":0.001504551,"about_ca_system_score_gemma":0.0059745926,"threshold_uncertainty_score":0.35446084},"labels":[],"label_agreement":null},{"id":"W6902027728","doi":"10.6084/m9.figshare.19759380","title":"Additional file 2 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test)","score_opus":0.06545233424165164,"score_gpt":0.4051565814282834,"score_spread":0.3397042471866318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6902027728","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002811582,0.0000147255105,0.00095326535,0.00010883958,0.000022619837,0.00008766399,0.9964204,0.001310295,0.0008011109],"genre_scores_gemma":[0.006328249,0.0000789042,0.009770464,0.0003094668,0.00009348639,0.0011588796,0.97611994,0.0012186328,0.004921907],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99911124,0.00013378268,0.00022478885,0.00024366229,0.00020800358,0.00007865346],"domain_scores_gemma":[0.9826717,0.012227457,0.0012287757,0.0009833799,0.002499627,0.00038921068],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015967397,0.0010129962,0.000758618,0.003278639,0.000535015,0.0014578437,0.0014361892,0.0008755269,0.6593007],"category_scores_gemma":[0.01875455,0.00044147318,0.0008404264,0.0035647932,0.00023829895,0.0013120739,0.0012530104,0.0007234714,0.11370514],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003630374,0.00009297599,0.0027271956,0.0023424171,0.00005830915,0.0001381781,0.000075683085,0.0004647164,0.00059998134,0.0007379249,0.9737857,0.018613936],"study_design_scores_gemma":[0.0021158403,0.00020999773,0.024722127,0.0013534038,0.0002107716,0.0007118265,0.0003476576,0.00457367,0.005519207,0.008421266,0.95165807,0.00015606242],"about_ca_topic_score_codex":0.0036780739,"about_ca_topic_score_gemma":0.0059268624,"teacher_disagreement_score":0.6593007,"about_ca_system_score_codex":0.0009827464,"about_ca_system_score_gemma":0.001942426,"threshold_uncertainty_score":0.4859662},"labels":[],"label_agreement":null},{"id":"W6902057674","doi":"10.6084/m9.figshare.19759383.v1","title":"Additional file 3 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test)","score_opus":0.06651094892406054,"score_gpt":0.4052983664027775,"score_spread":0.338787417478717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6902057674","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002848359,0.000014343533,0.0009517929,0.00010855453,0.000023196328,0.00009266124,0.9961677,0.0014929208,0.00086392136],"genre_scores_gemma":[0.0061974907,0.000077588564,0.009313242,0.00031999193,0.00009056983,0.0011563551,0.9763431,0.0013620872,0.0051396275],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990785,0.00013410773,0.00024003493,0.00024087721,0.00022203497,0.000084447944],"domain_scores_gemma":[0.98294026,0.011836558,0.0012496228,0.0009770559,0.0026081125,0.00038838587],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001610309,0.0010803692,0.0007746469,0.003381597,0.0005663586,0.0015177715,0.0014765777,0.0009171269,0.6729432],"category_scores_gemma":[0.018463261,0.00045347246,0.00090835465,0.0035535928,0.00025361346,0.0013500636,0.001313271,0.0007326297,0.1180505],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038950273,0.00009205604,0.0028794166,0.002437359,0.00006351767,0.00014936471,0.00008014758,0.00049578055,0.00063574873,0.00074951525,0.9740831,0.017944569],"study_design_scores_gemma":[0.002199735,0.00020658971,0.02405914,0.0014742648,0.00021482444,0.00069758086,0.0003627542,0.0045486535,0.0059243334,0.008897787,0.95124906,0.00016525632],"about_ca_topic_score_codex":0.004108138,"about_ca_topic_score_gemma":0.0065211216,"teacher_disagreement_score":0.6729432,"about_ca_system_score_codex":0.0010615097,"about_ca_system_score_gemma":0.0020312287,"threshold_uncertainty_score":0.46650684},"labels":[],"label_agreement":null},{"id":"W6902109808","doi":"10.6084/m9.figshare.20175986","title":"Additional file 6 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07351232016756304,"score_gpt":0.3383119798651275,"score_spread":0.26479965969756447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6902109808","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00012579093,0.000019057341,0.0011052993,0.00014212106,0.00005683083,0.000056591725,0.9941446,0.0018084724,0.0025412722],"genre_scores_gemma":[0.0043135225,0.00018426261,0.012015151,0.00044969222,0.000086063854,0.0007275076,0.96800214,0.004564379,0.009657185],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99945134,0.00008377704,0.00009438495,0.00013902257,0.00014218022,0.00008934551],"domain_scores_gemma":[0.9895538,0.007353454,0.00047066007,0.00074381905,0.001509138,0.00036909466],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016709118,0.0011292972,0.0010473955,0.003704612,0.00082991185,0.002490701,0.0018611593,0.0012802711,0.81457025],"category_scores_gemma":[0.01572349,0.000748244,0.0014034029,0.005031253,0.00045605673,0.0034576845,0.001768972,0.0014673845,0.26104245],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017135222,0.000052610514,0.00064526853,0.0015329907,0.000033017273,0.000058922327,0.00009097645,0.00031326155,0.00022958628,0.0021470312,0.9866666,0.008058523],"study_design_scores_gemma":[0.0010140722,0.000037135615,0.0036244614,0.000821176,0.00007511889,0.00022247696,0.00023813701,0.0011731541,0.0011655509,0.013038022,0.9785091,0.00008145008],"about_ca_topic_score_codex":0.008592859,"about_ca_topic_score_gemma":0.011564062,"teacher_disagreement_score":0.81457025,"about_ca_system_score_codex":0.0015505982,"about_ca_system_score_gemma":0.002367243,"threshold_uncertainty_score":0.26449305},"labels":[],"label_agreement":null},{"id":"W6902369799","doi":"10.6084/m9.figshare.20175992","title":"Additional file 8 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07353081139614696,"score_gpt":0.3383384590391222,"score_spread":0.2648076476429752,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6902369799","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00012953122,0.000018929011,0.0011650198,0.00014841625,0.000059230406,0.00005724887,0.9935835,0.0021005643,0.0027375487],"genre_scores_gemma":[0.004199554,0.00017987289,0.011825903,0.00043715595,0.000087438784,0.00072931004,0.96710634,0.0050149034,0.0104194805],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994454,0.00008405103,0.00009373432,0.00013443567,0.00015035218,0.000091965434],"domain_scores_gemma":[0.98938924,0.0074274307,0.00046097895,0.00075005664,0.0016039999,0.00036826913],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001674601,0.0011555204,0.0010252123,0.0037335786,0.00083065475,0.0024777704,0.0019038817,0.0012470192,0.8184529],"category_scores_gemma":[0.01613095,0.00075321534,0.0013778612,0.0049515697,0.0004540633,0.0035301987,0.001812802,0.0014505866,0.27913663],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016579278,0.000047990543,0.0005927556,0.0013538493,0.00002926609,0.000056446843,0.00008393448,0.00028177057,0.0002163738,0.0019500136,0.9871818,0.00804004],"study_design_scores_gemma":[0.0009790191,0.00003621916,0.0035230536,0.00084244367,0.00006954367,0.00021688618,0.00023945536,0.0011612254,0.001192128,0.012818412,0.9788432,0.000078439356],"about_ca_topic_score_codex":0.008144538,"about_ca_topic_score_gemma":0.011144303,"teacher_disagreement_score":0.8184529,"about_ca_system_score_codex":0.0015243109,"about_ca_system_score_gemma":0.0023057372,"threshold_uncertainty_score":0.25895494},"labels":[],"label_agreement":null},{"id":"W6904626135","doi":"10.1371/journal.pone.0284954.s002","title":"Complete list of terms used for database search.","year":2023,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Schizophrenia (object-oriented programming); Odds ratio; Subgroup analysis; Meta-analysis; Bonferroni correction; Depression (economics); Confidence interval; Web of science","score_opus":0.12891004808470197,"score_gpt":0.34915261475162,"score_spread":0.22024256666691805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6904626135","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010072975,0.099684075,0.001544349,0.0023095706,0.0009872213,0.012699472,0.859191,0.0007684389,0.021808574],"genre_scores_gemma":[0.019839283,0.3556466,0.024115713,0.012645015,0.0012988866,0.142844,0.40067542,0.0018285227,0.041106448],"study_design_codex":"systematic_review","study_design_gemma":"not_applicable","domain_scores_codex":[0.9953981,0.0007859755,0.0024963124,0.0005898111,0.0004836135,0.0002462097],"domain_scores_gemma":[0.9888911,0.006317794,0.0023779646,0.00023621124,0.0017785791,0.00039834448],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0032596225,0.0025430322,0.009965248,0.030274432,0.0009647783,0.0032055597,0.003850583,0.0021684852,0.3728585],"category_scores_gemma":[0.025348812,0.00091261114,0.004537142,0.043409083,0.00077579747,0.004464981,0.0027677917,0.0021063415,0.061816692],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00045900428,0.00004147802,0.00057821424,0.7411322,0.00071934337,0.00026762186,0.00033427682,0.0002471194,0.0002847717,0.0021463712,0.2103486,0.043440923],"study_design_scores_gemma":[0.0012093104,0.0002096431,0.0035666307,0.41196665,0.002890601,0.0008666954,0.00073814613,0.00038004172,0.00036502865,0.0038753904,0.5737674,0.0001644179],"about_ca_topic_score_codex":0.010374098,"about_ca_topic_score_gemma":0.013053137,"teacher_disagreement_score":0.3728585,"about_ca_system_score_codex":0.006198768,"about_ca_system_score_gemma":0.012661317,"threshold_uncertainty_score":0.89454126},"labels":[],"label_agreement":null},{"id":"W6905907922","doi":"10.15468/dl.zhapgv","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6905907922","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008991342,0.00004594264,0.00007534755,0.000052789164,0.000014434686,0.000008679742,0.99807113,0.00075297913,0.00088890735],"genre_scores_gemma":[0.00018910969,0.000039152368,0.0002777925,0.00004661362,0.0000027641831,0.000036807232,0.9988427,0.000138132,0.0004269677],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989471,0.00013915348,0.00014521845,0.00036395105,0.00024958994,0.00015511615],"domain_scores_gemma":[0.997957,0.00054628344,0.00019794739,0.0005351567,0.00049949286,0.00026411228],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.000926572,0.002392026,0.0014675957,0.0047387904,0.0010426423,0.0023334136,0.0030274938,0.0020448996,0.097644664],"category_scores_gemma":[0.0053620576,0.0008890807,0.001259762,0.008667111,0.0004709566,0.0022000012,0.0024713872,0.0019162784,0.14874518],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004216527,0.000017742315,0.00046211877,0.0005546277,0.000016648075,0.000023117951,0.000026407975,0.0001790904,0.00016506607,0.00048998685,0.99613744,0.0018855642],"study_design_scores_gemma":[0.00008844464,0.000011721398,0.0018381344,0.00016769412,0.00001556084,0.00006272598,0.00007582153,0.0003005737,0.0002841346,0.00093539193,0.99620026,0.000019661778],"about_ca_topic_score_codex":0.024930798,"about_ca_topic_score_gemma":0.044289436,"teacher_disagreement_score":0.9023553,"about_ca_system_score_codex":0.001776493,"about_ca_system_score_gemma":0.0026390974,"threshold_uncertainty_score":0.3266539},"labels":[],"label_agreement":null},{"id":"W6906108847","doi":"10.15468/dl.wa4842","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6906108847","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008683959,0.00004605423,0.000073976254,0.000052751002,0.000014450205,0.00000857385,0.99808705,0.00077255553,0.00085771177],"genre_scores_gemma":[0.00018862518,0.000039079663,0.00028212377,0.000044862823,0.0000028537313,0.000038635168,0.99882466,0.00014371156,0.00043550832],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989176,0.00014734827,0.00014645807,0.00037162114,0.0002607419,0.00015623812],"domain_scores_gemma":[0.99784434,0.0006057355,0.00020266567,0.00055516604,0.0005231013,0.0002689117],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009721238,0.0024740056,0.001490561,0.004928096,0.0010521349,0.0024243896,0.0030278338,0.002090979,0.101007946],"category_scores_gemma":[0.0056761242,0.00089457387,0.0012832034,0.009026261,0.0004697449,0.0021590889,0.0024704377,0.0019449744,0.15298109],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004023703,0.000017530061,0.0004402312,0.0005433572,0.000016789216,0.000022688202,0.00002607387,0.00018459922,0.0001494169,0.00046688798,0.99624014,0.0018520697],"study_design_scores_gemma":[0.00008912409,0.000011712331,0.0018374318,0.0001734524,0.000015827995,0.000061527215,0.00007578643,0.00032044938,0.00028497903,0.0009571943,0.99615234,0.000020168041],"about_ca_topic_score_codex":0.025245257,"about_ca_topic_score_gemma":0.043126214,"teacher_disagreement_score":0.89899206,"about_ca_system_score_codex":0.0017990805,"about_ca_system_score_gemma":0.0026284233,"threshold_uncertainty_score":0.33790523},"labels":[],"label_agreement":null},{"id":"W6906229781","doi":"10.17605/osf.io/ravpk","title":"Knowledge Synthesis Protocol Template","year":2023,"lang":"en","type":"article","venue":"OSF Preprints (OSF Preprints)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Protocol (science); Knowledge base; Knowledge-based systems; Knowledge representation and reasoning; Open Knowledge Base Connectivity","score_opus":0.02855859785865795,"score_gpt":0.31728537874996576,"score_spread":0.2887267808913078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6906229781","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016680724,0.0007585066,0.43827885,0.004251044,0.0021973362,0.03372079,0.37764984,0.045439553,0.096036054],"genre_scores_gemma":[0.01014734,0.001778037,0.3946385,0.0042519723,0.00038886172,0.09679251,0.40305656,0.014935509,0.074010655],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9847898,0.0057150894,0.004320475,0.0016994054,0.0028320774,0.00064306066],"domain_scores_gemma":[0.9297132,0.03828988,0.002252703,0.012589723,0.015633833,0.0015205602],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.030693615,0.0018825361,0.0016685883,0.0066174306,0.002255214,0.0054714535,0.002932062,0.0030608024,0.23827063],"category_scores_gemma":[0.06918291,0.0024593673,0.002003458,0.00498506,0.0016565678,0.0046581067,0.00517607,0.0041048387,0.13845342],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012246255,0.00024275604,0.00057435787,0.007982375,0.00011842071,0.00080823514,0.0020295752,0.0013619259,0.01231955,0.071057685,0.76422596,0.13805448],"study_design_scores_gemma":[0.00013502095,0.000042675016,0.00018909863,0.0008253202,0.0000472309,0.00013275903,0.00031380175,0.0007674393,0.007533483,0.012916381,0.9770423,0.000054415497],"about_ca_topic_score_codex":0.0071904305,"about_ca_topic_score_gemma":0.005494502,"teacher_disagreement_score":0.9693064,"about_ca_system_score_codex":0.003899839,"about_ca_system_score_gemma":0.016101174,"threshold_uncertainty_score":0.79709464},"labels":[],"label_agreement":null},{"id":"W6906393256","doi":"10.17605/osf.io/4k69q","title":"Comparing the National Library of Medicine (NLM)’s Medical Text Indexer (MTI) to Human Indexing: A Pilot Study","year":2022,"lang":"en","type":"article","venue":"OSF Preprints (OSF Preprints)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"National library; Medical library; MEDLINE; Medical journal; Digital library","score_opus":0.04820279487776538,"score_gpt":0.3148710831137579,"score_spread":0.26666828823599253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6906393256","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"evaluation","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"evaluation","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95163023,0.0021257452,0.013661893,0.0026876435,0.00035887232,0.002461615,0.014476293,0.0022228989,0.0103748515],"genre_scores_gemma":[0.93498814,0.00090974703,0.038932238,0.0007397937,0.0002307537,0.0013895169,0.019582417,0.0006928492,0.0025345946],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9739667,0.015278191,0.003459018,0.002575731,0.004144143,0.0005760881],"domain_scores_gemma":[0.76645845,0.18272422,0.007414387,0.017226856,0.021938499,0.004237567],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04836595,0.0006180813,0.0013075679,0.00596666,0.0012523463,0.0035094419,0.0019758483,0.0015115221,0.0071666213],"category_scores_gemma":[0.17115387,0.00033943477,0.0013570447,0.006307944,0.0014109113,0.007877285,0.003258244,0.0013837874,0.002280759],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0701324,0.016989743,0.2789791,0.008648978,0.0038126484,0.00063868635,0.010802571,0.0070761796,0.020265939,0.0060214126,0.03408956,0.5425429],"study_design_scores_gemma":[0.0117080035,0.04665317,0.68806,0.0012700153,0.008387591,0.002890165,0.013858948,0.0800965,0.042571437,0.01059659,0.09320499,0.00070251885],"about_ca_topic_score_codex":0.011347604,"about_ca_topic_score_gemma":0.0077686734,"teacher_disagreement_score":0.95163405,"about_ca_system_score_codex":0.0025532674,"about_ca_system_score_gemma":0.002938712,"threshold_uncertainty_score":0.25578666},"labels":[{"model":"gemma","categories":["metaresearch"],"domain":"methods","study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"observational","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W6908642897","doi":"10.3389/fninf.2018.00085.s002","title":"Presentation_2_National Neuroinformatics Framework for Canadian Consortium on Neurodegeneration in Aging (CCNA).PPTX","year":2018,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Neuroinformatics; Flagging; Upload; Data quality; Modular design; Quality (philosophy); Disease; Protocol (science)","score_opus":0.05142510241751563,"score_gpt":0.3277019302003282,"score_spread":0.27627682778281254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6908642897","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00072086527,0.0010936918,0.02839776,0.04613959,0.0070847576,0.002261493,0.35979182,0.0357144,0.51879567],"genre_scores_gemma":[0.009152849,0.0018310562,0.049483683,0.011865239,0.0016554175,0.002247517,0.33726102,0.01845768,0.5680455],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99211395,0.0009923887,0.00028517682,0.00064373785,0.004640335,0.0013245326],"domain_scores_gemma":[0.967505,0.002457636,0.00046842062,0.002661085,0.020136813,0.0067710825],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0137951635,0.0015776347,0.0011577493,0.003687979,0.004615205,0.012445697,0.0075476947,0.0042033615,0.514872],"category_scores_gemma":[0.02724924,0.0012783234,0.0014835607,0.0045634015,0.001543309,0.005501497,0.010335129,0.0037946973,0.280006],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019692423,0.00000992293,0.00010527029,0.000047279107,0.0000032159019,0.000018104336,0.000059695012,0.00006814999,0.00008465468,0.00318717,0.9860968,0.010299944],"study_design_scores_gemma":[0.000020303634,0.0000048688626,0.00040775555,0.00009582113,0.0000035557757,0.000019167865,0.00008704709,0.00018087098,0.00015533325,0.0020647675,0.99694127,0.00001928663],"about_ca_topic_score_codex":0.60375714,"about_ca_topic_score_gemma":0.63674796,"teacher_disagreement_score":0.9794434,"about_ca_system_score_codex":0.020556604,"about_ca_system_score_gemma":0.07374964,"threshold_uncertainty_score":0.7971528},"labels":[],"label_agreement":null},{"id":"W6912141488","doi":"10.5281/zenodo.3398964","title":"IMIA LaMB WG event: 'Biomedical Semantics in the Big Data Era', Workshop at MEDINFO 2015 – São Paulo,Brazil","year":2019,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Big data; Semantics (computer science); Presentation (obstetrics); Bridging (networking); Meaning (existential); Set (abstract data type)","score_opus":0.06545531387725695,"score_gpt":0.3302578687744909,"score_spread":0.264802554897234,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912141488","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.018179603,0.050756272,0.21597593,0.33499557,0.05175702,0.0022706222,0.07023716,0.018260272,0.23756756],"genre_scores_gemma":[0.10399571,0.041615143,0.13074468,0.026518224,0.023223381,0.003063089,0.12538244,0.01896259,0.5264946],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970837,0.0012201619,0.000157283,0.00044638137,0.00079009694,0.00030237768],"domain_scores_gemma":[0.99245864,0.0026462895,0.00023440455,0.000982696,0.001877134,0.0018008887],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014539889,0.0011645317,0.0011180014,0.0024618183,0.001696849,0.0056469105,0.0019071874,0.0025583715,0.052649032],"category_scores_gemma":[0.014083264,0.00059678796,0.0011572554,0.002056603,0.001462775,0.0065297247,0.007479958,0.002817906,0.02239073],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012892712,0.00003732628,0.00035259168,0.00033836765,0.0000187857,0.00012966015,0.00082536024,0.0004475691,0.0020005764,0.013896483,0.93242645,0.049397953],"study_design_scores_gemma":[0.0000166427,0.000021469807,0.00096880924,0.00029910833,0.000007805642,0.00011791155,0.0005107852,0.0008446482,0.0012335218,0.00661304,0.9893496,0.000016787842],"about_ca_topic_score_codex":0.0089321695,"about_ca_topic_score_gemma":0.009684913,"teacher_disagreement_score":0.052649032,"about_ca_system_score_codex":0.00339602,"about_ca_system_score_gemma":0.005283639,"threshold_uncertainty_score":0.17612857},"labels":[],"label_agreement":null},{"id":"W6912317890","doi":"10.5281/zenodo.16815811","title":"DEPRECATED_Matcher (Version 2) for Automated Task Alignment in the Genomic API for Model Evaluation (GAME)","year":2025,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Python (programming language); Workflow; Task (project management); Process (computing); Graph; Core (optical fiber); Architecture; Ontology alignment","score_opus":0.03450288535052127,"score_gpt":0.2927991912329317,"score_spread":0.2582963058824104,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912317890","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0025272751,0.00022551902,0.21023288,0.00021470715,0.00022966957,0.00024243417,0.025977131,0.75037473,0.009975693],"genre_scores_gemma":[0.07198035,0.00041921614,0.41928262,0.0011325601,0.00010040776,0.0019789427,0.1826849,0.28927752,0.033143476],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9984864,0.00024186641,0.0001175986,0.00045115527,0.0005022339,0.0002008249],"domain_scores_gemma":[0.99903595,0.00028394704,0.00006629371,0.00034140522,0.00019423747,0.000078152865],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0023717913,0.0026998252,0.0011442188,0.001357221,0.0008421005,0.0027931626,0.004254963,0.0015107724,0.11908343],"category_scores_gemma":[0.0066379425,0.0015568081,0.00216945,0.0010621202,0.0004994696,0.0031517851,0.0046907105,0.0026325602,0.064546645],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00085933163,0.00017516327,0.0028818175,0.00084722164,0.0002721497,0.00020556577,0.0003479757,0.009745163,0.0074805846,0.018599316,0.85495526,0.103630476],"study_design_scores_gemma":[0.00059957785,0.00020810924,0.0041400553,0.00028946612,0.00010631756,0.00034586806,0.00024552687,0.26374447,0.03723453,0.05060513,0.64218765,0.00029328465],"about_ca_topic_score_codex":0.009224968,"about_ca_topic_score_gemma":0.010414986,"teacher_disagreement_score":0.11908343,"about_ca_system_score_codex":0.0016276569,"about_ca_system_score_gemma":0.0019060267,"threshold_uncertainty_score":0.39837378},"labels":[],"label_agreement":null},{"id":"W6912383027","doi":"10.5281/zenodo.3665745","title":"CanDIG CHORD: Canadian Health Omics Repository, Distributed","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Software; Identification (biology); Software development; MEDLINE; Health data","score_opus":0.030924283322395836,"score_gpt":0.24732557746338868,"score_spread":0.21640129414099285,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912383027","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001002863,0.0017646982,0.021683812,0.0035340409,0.0007327607,0.00045533522,0.8443059,0.09247905,0.034041476],"genre_scores_gemma":[0.006043878,0.0021084757,0.042750634,0.0015026515,0.00019075304,0.0006385269,0.9051465,0.017185744,0.02443291],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974618,0.00021282428,0.0002624157,0.00040896528,0.001367004,0.00028701688],"domain_scores_gemma":[0.9835322,0.0023477552,0.0005034444,0.0035640027,0.007909602,0.002143088],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057980474,0.0026487568,0.0020290937,0.014168856,0.003738479,0.008021584,0.0073976372,0.0020295263,0.11688346],"category_scores_gemma":[0.024605935,0.0014801775,0.0016616272,0.021253044,0.0011314102,0.005099223,0.006203244,0.002672569,0.08228856],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019779059,0.000014474734,0.00059182325,0.00040290522,0.00005372131,0.000083205035,0.00009673747,0.00023062111,0.0005283104,0.0031092027,0.9657958,0.028895462],"study_design_scores_gemma":[0.00018346866,0.000011360619,0.002171806,0.00027522107,0.00007404537,0.000107040105,0.00011594862,0.0011094615,0.00095127383,0.0038062187,0.9911014,0.0000928312],"about_ca_topic_score_codex":0.63035935,"about_ca_topic_score_gemma":0.68255556,"teacher_disagreement_score":0.98856527,"about_ca_system_score_codex":0.011434719,"about_ca_system_score_gemma":0.056791063,"threshold_uncertainty_score":0.74363506},"labels":[],"label_agreement":null},{"id":"W6912459024","doi":"10.5281/zenodo.3381014","title":"IMIA LaMB WG event: 'Biomedical Semantics in the Big Data Era', Workshop at MEDINFO 2015 – São Paulo,Brazil","year":2019,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Thomson Reuters (Canada)","funders":"","keywords":"Big data; Semantics (computer science); Presentation (obstetrics); Bridging (networking); Meaning (existential); Set (abstract data type)","score_opus":0.05444336317088424,"score_gpt":0.30285898333210426,"score_spread":0.24841562016122,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912459024","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024023796,0.057866815,0.26853848,0.31322795,0.052068096,0.0023543711,0.055310782,0.015713388,0.21089637],"genre_scores_gemma":[0.1305943,0.04466768,0.15181163,0.025372455,0.026228298,0.003233956,0.11660039,0.016847955,0.48464334],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99701166,0.001262332,0.00016158604,0.00047184026,0.00076679984,0.00032567926],"domain_scores_gemma":[0.99337924,0.002241557,0.00022300567,0.00088284275,0.0016424656,0.0016308235],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.014389225,0.0012454594,0.0011684102,0.0022467128,0.0016236795,0.0052738716,0.0019151343,0.002621825,0.03942285],"category_scores_gemma":[0.012372284,0.00061321614,0.0011890138,0.0019263304,0.0014684106,0.0063715987,0.007212,0.0028278816,0.017012585],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019384082,0.000057742967,0.0004926479,0.00045245865,0.00002913876,0.0001871479,0.0010486144,0.00072865497,0.0031033794,0.019246904,0.90916824,0.065291226],"study_design_scores_gemma":[0.000025085234,0.000031997504,0.001257083,0.00034680372,0.000011685805,0.00017063068,0.0005754136,0.0013832233,0.0018857308,0.00923802,0.9850525,0.000021757463],"about_ca_topic_score_codex":0.007965092,"about_ca_topic_score_gemma":0.008120088,"teacher_disagreement_score":0.03942285,"about_ca_system_score_codex":0.0034329207,"about_ca_system_score_gemma":0.005257894,"threshold_uncertainty_score":0.13188255},"labels":[],"label_agreement":null},{"id":"W6912805906","doi":"10.5281/zenodo.5795448","title":"Access representation ontology developed for project's cohorts D3.5","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"European Commission","keywords":"Ontology; Data access; Data sharing; Process (computing); External Data Representation; Representation (politics); Controlled vocabulary; Data integration","score_opus":0.09342408288443421,"score_gpt":0.34731967819086695,"score_spread":0.25389559530643274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6912805906","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007162374,0.00052682427,0.698297,0.0036335606,0.0007048481,0.0015954622,0.20740482,0.033161625,0.047513425],"genre_scores_gemma":[0.03445254,0.0010188676,0.5031901,0.0017756415,0.00014634959,0.0022821254,0.4277974,0.0078733815,0.021463564],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99695146,0.0005328106,0.0005683369,0.0007012345,0.0010272529,0.00021887345],"domain_scores_gemma":[0.9962358,0.0008689516,0.00025616516,0.0011361402,0.0011536599,0.0003493514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005338648,0.00083861843,0.00070050644,0.004290587,0.0014045822,0.0032502783,0.0018121869,0.0013757056,0.016381059],"category_scores_gemma":[0.010100494,0.0007341799,0.0024341631,0.003145097,0.0007021731,0.004009358,0.0046918923,0.0026145459,0.0095148245],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041564798,0.00026288893,0.012669946,0.0024612793,0.00027077223,0.00077529065,0.0022465002,0.0059795887,0.010335207,0.3126064,0.43441537,0.21756105],"study_design_scores_gemma":[0.000035755646,0.000027111593,0.0028728838,0.00029066295,0.00006639614,0.00042430987,0.00026127612,0.005687426,0.0027131215,0.030987117,0.9565844,0.000049666043],"about_ca_topic_score_codex":0.029538756,"about_ca_topic_score_gemma":0.027754953,"teacher_disagreement_score":0.029538756,"about_ca_system_score_codex":0.0024921338,"about_ca_system_score_gemma":0.010037864,"threshold_uncertainty_score":0.058733642},"labels":[],"label_agreement":null},{"id":"W6920569761","doi":"10.6084/m9.figshare.20175959.v1","title":"Additional file 11 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07574097540738355,"score_gpt":0.34069025744204695,"score_spread":0.2649492820346634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920569761","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000120281176,0.000016746968,0.0009989525,0.00013739555,0.00005297823,0.000054611894,0.99426043,0.0018129918,0.002545533],"genre_scores_gemma":[0.0040266044,0.00016304046,0.010116097,0.00040682853,0.00007856628,0.0006957431,0.97046036,0.004391591,0.009661081],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994407,0.000087542816,0.00009265696,0.0001359261,0.0001508093,0.00009231764],"domain_scores_gemma":[0.9892603,0.007568349,0.0004861669,0.0007303727,0.0015805988,0.00037414],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016732567,0.0011551693,0.0010194086,0.0036622745,0.0008586453,0.0024855635,0.0018402536,0.0012402016,0.8182728],"category_scores_gemma":[0.016357902,0.00073136576,0.0013547912,0.0052901115,0.0004521415,0.0034810947,0.0018325826,0.0013894307,0.2743302],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016389645,0.0000470803,0.00060516153,0.001303879,0.00002824644,0.000056625115,0.000082601626,0.0002727967,0.00020488836,0.0019330092,0.9876916,0.0076100905],"study_design_scores_gemma":[0.0009716824,0.000035420544,0.0037255061,0.0007864697,0.00006592624,0.00020828516,0.00024446542,0.0010535556,0.0011480019,0.012013037,0.97966945,0.000078237244],"about_ca_topic_score_codex":0.008647431,"about_ca_topic_score_gemma":0.0115243355,"teacher_disagreement_score":0.8182728,"about_ca_system_score_codex":0.0016069448,"about_ca_system_score_gemma":0.0023406914,"threshold_uncertainty_score":0.25921166},"labels":[],"label_agreement":null},{"id":"W6920663416","doi":"10.6084/m9.figshare.20175995.v1","title":"Additional file 9 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.0735364846307192,"score_gpt":0.3383365658065004,"score_spread":0.2648000811757812,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920663416","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00012046335,0.00001720711,0.0010074008,0.0001371285,0.000056309258,0.00005278679,0.994013,0.0020301081,0.0025656032],"genre_scores_gemma":[0.0038636609,0.00016479993,0.010049065,0.00040540722,0.00008303874,0.00066795066,0.9703846,0.004825985,0.009555519],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994168,0.00008582642,0.000101606944,0.00014173816,0.0001571291,0.000096904296],"domain_scores_gemma":[0.98918253,0.007573526,0.0004857894,0.0007638493,0.0016170582,0.0003772739],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016581252,0.0011948637,0.0010537851,0.0037710092,0.00085526233,0.0026227916,0.0018994344,0.0012836031,0.81717783],"category_scores_gemma":[0.016553968,0.00076930504,0.0013971413,0.0050571137,0.00046624165,0.0036626172,0.0019066909,0.0014599346,0.28215724],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017196896,0.000047223635,0.00060931814,0.0013290436,0.000030104638,0.000059324313,0.000082800274,0.00026580086,0.00021536421,0.0019281452,0.9874228,0.007838278],"study_design_scores_gemma":[0.0010064173,0.000035586254,0.0036349834,0.0008342868,0.000068556634,0.00021287991,0.00023588164,0.0010700687,0.0011812487,0.012493048,0.97914827,0.00007869949],"about_ca_topic_score_codex":0.008403507,"about_ca_topic_score_gemma":0.011176898,"teacher_disagreement_score":0.81717783,"about_ca_system_score_codex":0.0015732386,"about_ca_system_score_gemma":0.0023093587,"threshold_uncertainty_score":0.26077366},"labels":[],"label_agreement":null},{"id":"W6920682026","doi":"10.6084/m9.figshare.25389419","title":"Additional file 1 of Text analysis framework for identifying mutations among non-small cell lung cancer patients from laboratory data","year":2024,"lang":"en","type":"article","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Lung cancer; Mutation; Text mining; Patient data; Cancer","score_opus":0.036735448579721296,"score_gpt":0.34342847240375113,"score_spread":0.30669302382402985,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920682026","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00031048286,0.0000210546,0.0018107255,0.0001055702,0.000025389674,0.00008311349,0.9950506,0.001889671,0.0007033756],"genre_scores_gemma":[0.007622766,0.00009203082,0.018836116,0.00031685736,0.00007196724,0.0010378222,0.9667145,0.0015236344,0.003784374],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99936825,0.00010713651,0.00012961059,0.00019328632,0.00014093926,0.00006084095],"domain_scores_gemma":[0.98800987,0.009359418,0.00048622309,0.00052303803,0.0012901527,0.00033132092],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001443648,0.0013214326,0.0007409987,0.0031297987,0.000544172,0.0015228571,0.0014270004,0.0011129719,0.61132467],"category_scores_gemma":[0.017564915,0.00043736256,0.0008666269,0.0026481396,0.00030291738,0.0014395171,0.0013474937,0.00087402615,0.1032877],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041369878,0.00013910612,0.0038272964,0.0037653083,0.000084804116,0.0002298401,0.00013254979,0.0010339511,0.00096222805,0.0015507276,0.9605948,0.027265605],"study_design_scores_gemma":[0.001797853,0.0002259498,0.020139182,0.0016788938,0.00023045724,0.0012687121,0.00045952515,0.011268119,0.005563127,0.020992741,0.9362087,0.00016676517],"about_ca_topic_score_codex":0.003432529,"about_ca_topic_score_gemma":0.0068296352,"teacher_disagreement_score":0.61132467,"about_ca_system_score_codex":0.00091035693,"about_ca_system_score_gemma":0.0016139607,"threshold_uncertainty_score":0.5543982},"labels":[],"label_agreement":null},{"id":"W6920724989","doi":"10.6084/m9.figshare.12799365.v2","title":"Additional file 3 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2021,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Documentation; Rare events","score_opus":0.05121493344011546,"score_gpt":0.32572751215804435,"score_spread":0.2745125787179289,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920724989","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008901042,0.00007236454,0.00070557254,0.00020745854,0.00003019392,0.0005065889,0.9970655,0.00038425523,0.0009390778],"genre_scores_gemma":[0.0045947437,0.00077534886,0.029339358,0.0019480088,0.0002107669,0.023585193,0.9258939,0.0018402211,0.011812465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99626476,0.0009516968,0.0014146002,0.0006627695,0.00048447694,0.00022170082],"domain_scores_gemma":[0.8921012,0.09094481,0.0062795654,0.003198,0.006246082,0.0012303694],"candidate_categories":["metaresearch","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0081131775,0.0012668817,0.001841197,0.008411752,0.00089837535,0.0026006522,0.0020116447,0.0016588287,0.77455777],"category_scores_gemma":[0.08511096,0.0009813752,0.002512411,0.009701538,0.0005092184,0.002585809,0.00239626,0.0014314831,0.09563002],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042649359,0.000046355544,0.0013427264,0.04466585,0.00028185258,0.00007542654,0.00016752457,0.0004386585,0.00024391935,0.002219524,0.93413883,0.015952915],"study_design_scores_gemma":[0.003629256,0.00012571755,0.0070801587,0.016789561,0.00081847137,0.00025288577,0.00030437156,0.0006466138,0.00079251,0.014206558,0.955212,0.00014194199],"about_ca_topic_score_codex":0.004817988,"about_ca_topic_score_gemma":0.013929324,"teacher_disagreement_score":0.99798834,"about_ca_system_score_codex":0.0016435151,"about_ca_system_score_gemma":0.0063235867,"threshold_uncertainty_score":0.32156593},"labels":[],"label_agreement":null},{"id":"W6920741584","doi":"10.6084/m9.figshare.12228167","title":"Additional file 1 of The PRECISE (PREgnancy Care Integrating translational Science, Everywhere) database: open-access data collection in maternal and newborn health","year":2020,"lang":"en","type":"article","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Data collection; Health care; Health data; MEDLINE; Patient data","score_opus":0.10407266696264492,"score_gpt":0.38857583246429084,"score_spread":0.2845031655016459,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920741584","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007598581,0.000015892305,0.0001756259,0.000059480455,0.000010662651,0.000039857005,0.99901867,0.00014749505,0.00045635094],"genre_scores_gemma":[0.003677363,0.00016183307,0.0034410374,0.00039696036,0.00006403543,0.0012145637,0.98750883,0.0005817476,0.0029536171],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986958,0.0002719102,0.00033541542,0.0003033091,0.0002512488,0.0001423847],"domain_scores_gemma":[0.9705377,0.023028964,0.0017000253,0.0011714596,0.0026460423,0.00091579],"candidate_categories":["open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.002632321,0.0010049121,0.0010865984,0.0037404404,0.0006845167,0.001918662,0.0017374353,0.0011641103,0.70650226],"category_scores_gemma":[0.034821928,0.000572558,0.000803722,0.006588498,0.00037993706,0.0018226105,0.0017468185,0.001133461,0.11820088],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024984032,0.000057428813,0.002606515,0.0032535484,0.000050438506,0.000048138216,0.000083340456,0.0002467634,0.000114111106,0.0009353354,0.98434883,0.008005748],"study_design_scores_gemma":[0.0019321289,0.00011411389,0.020131806,0.00303528,0.0002058321,0.00036662442,0.0005635694,0.00081889465,0.0009502779,0.010360095,0.96139294,0.00012844383],"about_ca_topic_score_codex":0.00824011,"about_ca_topic_score_gemma":0.013763833,"teacher_disagreement_score":0.9982626,"about_ca_system_score_codex":0.0012824129,"about_ca_system_score_gemma":0.002643254,"threshold_uncertainty_score":0.41863889},"labels":[],"label_agreement":null},{"id":"W6920761972","doi":"10.6084/m9.figshare.12799365.v1","title":"Additional file 3 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2020,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Documentation; Rare events","score_opus":0.06675844190939934,"score_gpt":0.3214760123045299,"score_spread":0.25471757039513054,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920761972","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000087549786,0.000072269955,0.0006940231,0.00020691939,0.000030006055,0.00050526933,0.9970849,0.00037904974,0.0009399872],"genre_scores_gemma":[0.0045686495,0.0007795581,0.02892815,0.001961967,0.0002121664,0.023636252,0.92621857,0.0018254347,0.01186929],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99626595,0.00095160585,0.001412917,0.0006604106,0.0004866686,0.00022242343],"domain_scores_gemma":[0.8920568,0.090967745,0.006286347,0.0032144804,0.006240954,0.0012335872],"candidate_categories":["metaresearch","open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008116647,0.0012645345,0.0018421805,0.008431336,0.0008999852,0.0025966256,0.002021678,0.0016694728,0.77747566],"category_scores_gemma":[0.085223086,0.0009830914,0.0025106764,0.0097416,0.0005075872,0.0026026927,0.0024063943,0.0014339684,0.096612826],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042467698,0.00004612008,0.0013221144,0.04443338,0.00027916676,0.000074496995,0.0001656272,0.0004325193,0.00024033594,0.0022071025,0.9345319,0.015842475],"study_design_scores_gemma":[0.0036580486,0.00012622439,0.007083695,0.016926134,0.0008153395,0.00025379081,0.00030251604,0.00064093113,0.0007845722,0.014178331,0.95508814,0.000142218],"about_ca_topic_score_codex":0.00484238,"about_ca_topic_score_gemma":0.014018838,"teacher_disagreement_score":0.9979783,"about_ca_system_score_codex":0.0016452193,"about_ca_system_score_gemma":0.0063576126,"threshold_uncertainty_score":0.31740397},"labels":[],"label_agreement":null},{"id":"W6920793498","doi":"10.60692/qpsch-q1g26","title":"Future-proofing and maximizing the utility of metadata: The PHA4GE SARS-CoV-2 contextual data specification package","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Dalhousie University; McMaster University; BC Centre for Disease Control; Simon Fraser University","funders":"","keywords":"Interoperability; Consistency (knowledge bases); Standardization; Harmonization; Contextual design; Data integration; Data aggregator; Data quality; Data consistency; Alliance","score_opus":0.10648704050317989,"score_gpt":0.27160296818766494,"score_spread":0.16511592768448505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920793498","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053257,0.0004026599,0.8932614,0.0063521247,0.0005850911,0.0018623208,0.034880105,0.04214803,0.01518253],"genre_scores_gemma":[0.025218792,0.00080293027,0.8606451,0.0027043456,0.0002642512,0.0020355026,0.090431765,0.010413766,0.0074835303],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98052067,0.007575734,0.0039986656,0.001453796,0.0054073473,0.0010438622],"domain_scores_gemma":[0.95510644,0.012810058,0.0027479504,0.017408175,0.010161257,0.0017661541],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.040992353,0.0014732435,0.0009528674,0.0044257403,0.0017106859,0.006855598,0.0038054178,0.0024595452,0.009659569],"category_scores_gemma":[0.0544225,0.0016045116,0.002537274,0.0034584226,0.0019048713,0.007890446,0.008617221,0.0039484813,0.009558591],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008619076,0.0004322967,0.013782091,0.0018657506,0.00024255374,0.0010296183,0.0032097301,0.013029661,0.015808472,0.3257431,0.38810214,0.23589277],"study_design_scores_gemma":[0.00014738046,0.00012993076,0.0023974339,0.0012110089,0.00009336534,0.000521788,0.00070055,0.016972678,0.015130654,0.08299392,0.87947524,0.00022614063],"about_ca_topic_score_codex":0.013373684,"about_ca_topic_score_gemma":0.012236432,"teacher_disagreement_score":0.9590076,"about_ca_system_score_codex":0.002570522,"about_ca_system_score_gemma":0.012620243,"threshold_uncertainty_score":0.21679085},"labels":[],"label_agreement":null},{"id":"W6920873072","doi":"10.6084/m9.figshare.19759377.v1","title":"Additional file 1 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test); Medical record","score_opus":0.06772187932270748,"score_gpt":0.40586719390939824,"score_spread":0.3381453145866907,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920873072","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00027153597,0.000014436187,0.00096313044,0.00010285164,0.000020648316,0.00008745561,0.99660635,0.0012206004,0.0007130035],"genre_scores_gemma":[0.0061854185,0.00008205972,0.0104152225,0.00029521153,0.000091530004,0.0011791693,0.9762695,0.00112502,0.0043569333],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991085,0.00014093381,0.00022678038,0.00024385689,0.00020369023,0.00007623959],"domain_scores_gemma":[0.98120236,0.013487145,0.0012818915,0.0010366353,0.0025829251,0.00040915623],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016823882,0.0009833763,0.0007262373,0.003235851,0.00052174047,0.0013788228,0.0014584729,0.0008273894,0.6584403],"category_scores_gemma":[0.020382663,0.0004426881,0.00077036617,0.0036481943,0.0002325513,0.0013226898,0.0012704524,0.0007357723,0.11987488],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033828904,0.00009682073,0.0026232966,0.0022136716,0.00005458927,0.00011809142,0.00007055125,0.000464638,0.0005834695,0.00070051174,0.97309047,0.0196456],"study_design_scores_gemma":[0.002169057,0.00024597652,0.026303245,0.0013870009,0.00021544506,0.00075047126,0.00034244254,0.004969887,0.005920783,0.008623487,0.94891137,0.00016083439],"about_ca_topic_score_codex":0.0033767682,"about_ca_topic_score_gemma":0.005563368,"teacher_disagreement_score":0.6584403,"about_ca_system_score_codex":0.0009855683,"about_ca_system_score_gemma":0.0019882405,"threshold_uncertainty_score":0.48719347},"labels":[],"label_agreement":null},{"id":"W6920927648","doi":"10.6084/m9.figshare.27246017.v1","title":"Additional file 1 of Dynamic Retrieval Augmented Generation of Ontologies using Artificial Intelligence (DRAGON-AI)","year":2024,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Applications of artificial intelligence; Key (lock); Ontology; Field (mathematics); Expert system","score_opus":0.08765391331783524,"score_gpt":0.33007154077081746,"score_spread":0.24241762745298223,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920927648","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00017253605,0.000021658945,0.00086038135,0.000077467186,0.000046170375,0.000047405665,0.99530494,0.0024046588,0.0010647981],"genre_scores_gemma":[0.0038337344,0.0000789515,0.008700819,0.00026513688,0.000052195002,0.0006604935,0.97884136,0.0029881867,0.0045791254],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99923766,0.00015208997,0.00009915171,0.00022882284,0.00019934298,0.00008308718],"domain_scores_gemma":[0.9865726,0.010421902,0.00032090602,0.00096522935,0.0013299962,0.00038928966],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016827423,0.0013371144,0.0010862511,0.0025766178,0.00078242336,0.0025027357,0.0021246104,0.0014297625,0.796424],"category_scores_gemma":[0.020845756,0.0006832656,0.0009709775,0.0033879492,0.0003595789,0.0017168386,0.0018014953,0.001256539,0.2593513],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022696152,0.00005721486,0.00063644495,0.0014342837,0.000040645184,0.00006389545,0.000042884407,0.00047826397,0.00022569625,0.00075158506,0.98719656,0.008845637],"study_design_scores_gemma":[0.0020527348,0.00015963604,0.007031918,0.0009737337,0.00014951962,0.00037461577,0.00023008356,0.0045056967,0.0027818936,0.014255987,0.96733165,0.0001526387],"about_ca_topic_score_codex":0.0042391387,"about_ca_topic_score_gemma":0.00848256,"teacher_disagreement_score":0.796424,"about_ca_system_score_codex":0.0009849994,"about_ca_system_score_gemma":0.0014936058,"threshold_uncertainty_score":0.29037648},"labels":[],"label_agreement":null},{"id":"W6920961910","doi":"10.6084/m9.figshare.14684514","title":"Additional file 1 of Text mining to support abstract screening for knowledge syntheses: a semi-automated workflow","year":2021,"lang":"en","type":"article","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Queen's University; Toronto Metropolitan University","funders":"","keywords":"Workflow; Knowledge extraction; Text mining; Key (lock); File format; Web mining","score_opus":0.06495120301387569,"score_gpt":0.34190870977133186,"score_spread":0.2769575067574562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6920961910","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002712535,0.000032610275,0.003018657,0.00017435268,0.00006185134,0.00017940014,0.9903424,0.0040346524,0.0018848154],"genre_scores_gemma":[0.0052690976,0.00014183822,0.026421372,0.00054152886,0.00017043033,0.0018535499,0.95049673,0.0047468394,0.010358507],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9986614,0.0002436376,0.00031360035,0.00038003337,0.00030490826,0.00009645721],"domain_scores_gemma":[0.96700716,0.025598748,0.0014862224,0.0016960142,0.0034995081,0.00071233767],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0028431045,0.0016278581,0.0010685968,0.003921063,0.0008692846,0.0025777635,0.0020536322,0.001257885,0.7856534],"category_scores_gemma":[0.03292826,0.00073257595,0.001049658,0.0042629684,0.0004809393,0.0024582422,0.002542303,0.0012860488,0.24864016],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00041553646,0.00008180323,0.0013432179,0.0032137805,0.000054101438,0.00014435542,0.00011225173,0.00038988006,0.00061461166,0.0010659948,0.97080314,0.021761276],"study_design_scores_gemma":[0.0012300111,0.00012787902,0.007289323,0.0013508642,0.0001505947,0.00054130115,0.00031308006,0.0025485766,0.0034617414,0.015459293,0.9673921,0.00013522261],"about_ca_topic_score_codex":0.0029658023,"about_ca_topic_score_gemma":0.0050212615,"teacher_disagreement_score":0.9971569,"about_ca_system_score_codex":0.0011797571,"about_ca_system_score_gemma":0.0026047097,"threshold_uncertainty_score":0.30573934},"labels":[],"label_agreement":null},{"id":"W6921103668","doi":"10.6084/m9.figshare.21397896","title":"Additional file 3 of Machine learning algorithms to identify cluster randomized trials from MEDLINE and EMBASE","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Ottawa Hospital; McMaster University; London Health Sciences Centre; Lawson Health Research Institute","funders":"","keywords":"Cluster (spacecraft); Convolutional neural network; MEDLINE; Feature (linguistics); Key (lock)","score_opus":0.045879010219779,"score_gpt":0.32122997232150885,"score_spread":0.2753509621017298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6921103668","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00016080082,0.000040378258,0.00045926872,0.00009055395,0.00001806003,0.00012133289,0.9981382,0.00050532183,0.00046610867],"genre_scores_gemma":[0.009694142,0.0003307332,0.010657212,0.00067088223,0.00015186727,0.0047437893,0.9642343,0.0015790561,0.007938032],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99862695,0.00025833823,0.00038253065,0.0003468046,0.0002431543,0.00014221991],"domain_scores_gemma":[0.9507465,0.042238228,0.00236856,0.0014561297,0.0026429864,0.00054766063],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0029725288,0.0014591111,0.0018794797,0.0047663045,0.0006398561,0.0020365734,0.0020083894,0.0015724719,0.8534177],"category_scores_gemma":[0.046422873,0.0008275629,0.001840787,0.007201928,0.00040035127,0.0021957273,0.0013407909,0.0010629102,0.12418424],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001118065,0.00015526016,0.003133337,0.022153296,0.0003176727,0.00017178425,0.00008893514,0.0015960089,0.00034494032,0.0018407487,0.9484676,0.020612316],"study_design_scores_gemma":[0.019630747,0.00082114676,0.028544618,0.010246882,0.0011368118,0.0010309049,0.00039540188,0.008552859,0.0026070804,0.034428664,0.89225274,0.0003522643],"about_ca_topic_score_codex":0.0052995365,"about_ca_topic_score_gemma":0.0135017885,"teacher_disagreement_score":0.99702746,"about_ca_system_score_codex":0.0017595406,"about_ca_system_score_gemma":0.0028626309,"threshold_uncertainty_score":0.20908189},"labels":[],"label_agreement":null},{"id":"W6923073486","doi":"10.1371/journal.pone.0057614.g001","title":"Photographs of the four sympatric nesting species of geese studied in the investigation on the Yukon-Kuskokwim Delta, Alaska: A) Emperor goose (&lt;i&gt;Chen canagica&lt;/i&gt;), B) Greater white-fronted goose (&lt;i&gt;Anser albifrons&lt;/i&gt;), C) Cackling goose (&lt;i&gt;Branta hutchinsii&lt;/i&gt;), and D) Black brant (&lt;i&gt;Branta bernicla&lt;/i&gt;).","year":2015,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Goose; Sympatric speciation; Emperor; Nesting (process); Annals","score_opus":0.04023825845385966,"score_gpt":0.24155431804802144,"score_spread":0.20131605959416177,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6923073486","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1945276,0.0008986428,0.001756494,0.00053058576,0.00048528938,0.0007677829,0.18272255,0.000582904,0.6177282],"genre_scores_gemma":[0.53323555,0.001567847,0.009643757,0.00050705817,0.00020630866,0.0004944191,0.100918174,0.00037776591,0.35304913],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99995005,0.0000060032303,0.0000034909367,0.000012894738,0.000014090496,0.000013525959],"domain_scores_gemma":[0.99987686,0.000027457329,0.0000133362655,0.000011906912,0.00004652675,0.000023955698],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000072127536,0.0003725626,0.00013256528,0.0015889254,0.00086024794,0.00026685212,0.00037007142,0.0003140925,0.18137953],"category_scores_gemma":[0.00015850119,0.0002050673,0.0001783116,0.0018248752,0.00024483714,0.00035872348,0.0004242726,0.00031002588,0.020614754],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00051001704,0.00026114637,0.0823077,0.0009303007,0.00007019369,0.0018030843,0.0066321357,0.0013143531,0.025661012,0.002014081,0.73604614,0.14244998],"study_design_scores_gemma":[0.000020647209,0.000062545994,0.67089915,0.00014346553,0.000020969095,0.00071685197,0.004617512,0.0001777674,0.0005790468,0.00016352386,0.32258344,0.000015146456],"about_ca_topic_score_codex":0.05179253,"about_ca_topic_score_gemma":0.26357204,"teacher_disagreement_score":0.18137953,"about_ca_system_score_codex":0.00034507056,"about_ca_system_score_gemma":0.00021255817,"threshold_uncertainty_score":0.6067749},"labels":[],"label_agreement":null},{"id":"W6923218763","doi":"10.1371/journal.pone.0041468.g004","title":"Low-impact areas for wind development in Saskatchewan. (","year":2015,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ecoregion; Wildlife; Wind power; Maximum sustained wind","score_opus":0.03504953443136096,"score_gpt":0.30199476560676286,"score_spread":0.2669452311754019,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6923218763","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06084533,0.0023685037,0.002569353,0.005797688,0.0011731505,0.0012624643,0.7346491,0.0012978194,0.19003661],"genre_scores_gemma":[0.2775185,0.004825579,0.023547528,0.0049975114,0.00014100774,0.003045208,0.34344566,0.00074900466,0.34173],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997634,0.000035821115,0.00001647827,0.000036299127,0.0000518517,0.00009612951],"domain_scores_gemma":[0.9992623,0.00006754513,0.000045632165,0.00004762203,0.00034181293,0.00023515437],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00040233159,0.00059459696,0.00025700277,0.00095890596,0.0013534378,0.0014012767,0.00097509683,0.0005751421,0.10407524],"category_scores_gemma":[0.0009303232,0.0003634045,0.0004007946,0.002638258,0.0002873161,0.0006707793,0.0014162173,0.00069670525,0.014502305],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002748274,0.00011678761,0.07477759,0.0012030436,0.000116609546,0.0005812335,0.0007549967,0.0015433684,0.0031366125,0.004137419,0.7885459,0.12481161],"study_design_scores_gemma":[0.00024551334,0.000050940485,0.43178162,0.00088321185,0.0000633999,0.00018708895,0.0041055465,0.0014890342,0.0008835607,0.0013916761,0.5588208,0.00009762871],"about_ca_topic_score_codex":0.89276457,"about_ca_topic_score_gemma":0.9676033,"teacher_disagreement_score":0.10723543,"about_ca_system_score_codex":0.0064231306,"about_ca_system_score_gemma":0.030214414,"threshold_uncertainty_score":0.34816635},"labels":[],"label_agreement":null},{"id":"W6923878027","doi":"10.14288/1.0446444","title":"Canadian Development Company's steamer 'Columbian' and the Yukon Flyer Line Company's 'Eldorado' starting from Dawson July 4th 99 on a race to White Horse rapids","year":2024,"lang":"en","type":"other","venue":"Open Collections","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Race (biology); White (mutation); Line (geometry); Horse racing","score_opus":0.018684711619264424,"score_gpt":0.2716431784741865,"score_spread":0.25295846685492207,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6923878027","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0011446257,0.00059257826,0.0008262029,0.005152764,0.0006305798,0.000074146104,0.021746056,0.0022327441,0.9676003],"genre_scores_gemma":[0.0021866912,0.00042336804,0.00094732275,0.00045419764,0.000026318534,0.000014928258,0.007862436,0.00058621104,0.98749846],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991297,0.000025176232,0.000018851884,0.00010315337,0.000579468,0.00014365527],"domain_scores_gemma":[0.9980934,0.00009948261,0.000049546554,0.00013086098,0.0010686996,0.0005580043],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00075686444,0.00084537076,0.0002973849,0.0026184556,0.005078514,0.004985095,0.00062661566,0.0012399713,0.34144506],"category_scores_gemma":[0.0020491788,0.0003819615,0.00035828754,0.006440555,0.0014620472,0.0028309545,0.0016459135,0.0013395196,0.12943672],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011152662,0.0000050405997,0.00013662293,0.000024697103,7.283999e-7,0.00002310266,0.0000885396,0.000016492222,0.00018246233,0.0035458386,0.9808764,0.015089028],"study_design_scores_gemma":[0.0000014574027,0.0000010123767,0.0005118091,0.000011884538,6.192544e-7,0.000008189807,0.00009342521,0.000015262478,0.00008278199,0.00013608788,0.9991341,0.0000034515135],"about_ca_topic_score_codex":0.83453107,"about_ca_topic_score_gemma":0.94988936,"teacher_disagreement_score":0.34144506,"about_ca_system_score_codex":0.014455064,"about_ca_system_score_gemma":0.032791454,"threshold_uncertainty_score":0.9393487},"labels":[],"label_agreement":null},{"id":"W6924256624","doi":"10.15468/dl.k8q5bb","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics)","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6924256624","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008043528,0.000048342496,0.00005876628,0.0000507003,0.000012450786,0.000008543747,0.99827456,0.00062960916,0.0008364904],"genre_scores_gemma":[0.00017535946,0.000037068745,0.00019520232,0.000043392796,0.000002670286,0.00003200756,0.9990501,0.0001122173,0.00035198903],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988952,0.00015024105,0.00014741496,0.00036672078,0.0002659651,0.00017436148],"domain_scores_gemma":[0.99799097,0.0005260153,0.00019154159,0.00053524075,0.0004799222,0.0002763383],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010447437,0.0024047708,0.0015853846,0.005576445,0.00101975,0.0024915235,0.003043879,0.0022428199,0.10387062],"category_scores_gemma":[0.005192833,0.0008896205,0.0012238001,0.009434359,0.00048694856,0.0022406455,0.0026171335,0.0019709761,0.16649933],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000044673947,0.000017818127,0.0004163794,0.0006158676,0.000018874027,0.000025303565,0.000024874258,0.00016181152,0.00018689975,0.00045766414,0.99614316,0.0018867451],"study_design_scores_gemma":[0.00008966004,0.000010369929,0.0018767385,0.00018929945,0.000016273136,0.000060846407,0.000070328344,0.0002550207,0.00031695192,0.00087839155,0.99621767,0.000018454655],"about_ca_topic_score_codex":0.020514077,"about_ca_topic_score_gemma":0.033982288,"teacher_disagreement_score":0.89612937,"about_ca_system_score_codex":0.0018319302,"about_ca_system_score_gemma":0.0025132906,"threshold_uncertainty_score":0.34748185},"labels":[],"label_agreement":null},{"id":"W6924438589","doi":"10.15468/dl.9xrys2","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Identification (biology); State (computer science)","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6924438589","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000070004404,0.000040961375,0.00006997544,0.00005422398,0.000013519084,0.000007629474,0.9983358,0.00063524646,0.0007726511],"genre_scores_gemma":[0.00016546476,0.00003452703,0.00024539494,0.000043712284,0.0000025447362,0.000035198555,0.99897397,0.00012854164,0.0003706212],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988533,0.00014577813,0.00015116605,0.00039661973,0.00028592537,0.000167182],"domain_scores_gemma":[0.9978125,0.00058231334,0.00020142084,0.0006102044,0.0005167398,0.0002768829],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010727365,0.0024389592,0.0016765536,0.0048487675,0.0011680896,0.0026449931,0.003333607,0.0024317745,0.103988744],"category_scores_gemma":[0.005647117,0.0009917973,0.0013016367,0.008740421,0.0005146786,0.0024830063,0.0027529334,0.0022855448,0.16390537],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000037760547,0.000017360633,0.00043086035,0.0005548089,0.000017748785,0.00002302421,0.0000278415,0.00017711156,0.00016230273,0.00048459246,0.99634117,0.0017253426],"study_design_scores_gemma":[0.00008312028,0.000008953265,0.0018941774,0.00018392301,0.00001576112,0.000059940932,0.000082178776,0.0002754739,0.00028207054,0.00091661845,0.99617827,0.000019550618],"about_ca_topic_score_codex":0.023960898,"about_ca_topic_score_gemma":0.04056392,"teacher_disagreement_score":0.89601123,"about_ca_system_score_codex":0.0019745221,"about_ca_system_score_gemma":0.0026934312,"threshold_uncertainty_score":0.34787697},"labels":[],"label_agreement":null},{"id":"W6924510007","doi":"10.15468/dl.x6dcew","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Download; Range (aeronautics); State (computer science); Identification (biology)","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6924510007","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0000863411,0.000043410644,0.000070173264,0.000047693775,0.000014560379,0.000008477772,0.9982988,0.0005751486,0.000855353],"genre_scores_gemma":[0.00018099735,0.000036934776,0.0002502625,0.000044145316,0.000002944603,0.00004127896,0.99890494,0.00012349988,0.0004148918],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988908,0.00015566905,0.00014637556,0.00038971176,0.00025632614,0.00016108515],"domain_scores_gemma":[0.99786085,0.0006227913,0.00020446088,0.0005311643,0.000492891,0.0002877117],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010209384,0.002488044,0.0015131149,0.005513534,0.0010849945,0.0024174005,0.0030043027,0.0022450662,0.10677911],"category_scores_gemma":[0.0052490225,0.00092061184,0.0012401774,0.009538329,0.00050429435,0.0021882234,0.0025472594,0.0020400982,0.15620665],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004302799,0.000019996578,0.00047287683,0.0006359046,0.000018013565,0.000026235819,0.000027053791,0.0001823099,0.00018249221,0.00046872816,0.99604416,0.001879155],"study_design_scores_gemma":[0.00009291308,0.000011547146,0.0018783525,0.00017926552,0.00001642361,0.00006606103,0.00007521554,0.00027042147,0.00029526025,0.00088835,0.9962057,0.000020538266],"about_ca_topic_score_codex":0.020673191,"about_ca_topic_score_gemma":0.0359635,"teacher_disagreement_score":0.8932209,"about_ca_system_score_codex":0.0017107324,"about_ca_system_score_gemma":0.002520231,"threshold_uncertainty_score":0.3572117},"labels":[],"label_agreement":null},{"id":"W6924802006","doi":"10.15468/dl.tyxh73","title":"Occurrence Download","year":2022,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Download; Alien; Range (aeronautics); State (computer science)","score_opus":0.014527195957572651,"score_gpt":0.23482297486680234,"score_spread":0.2202957789092297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6924802006","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007328273,0.000040645555,0.0000623775,0.000047945723,0.000014065012,0.0000084583,0.99832815,0.0006087282,0.0008163999],"genre_scores_gemma":[0.00017466232,0.000037516435,0.00025876492,0.000046601253,0.0000030621495,0.000042936656,0.99887496,0.00013077028,0.0004308007],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99895346,0.00014016902,0.0001407179,0.00037527096,0.00023364366,0.00015662235],"domain_scores_gemma":[0.99794155,0.00058946747,0.00019169274,0.00051768986,0.00049481104,0.00026491165],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009698259,0.0025245252,0.0015353193,0.005003958,0.0010442825,0.0024126917,0.002973573,0.0022005024,0.11550637],"category_scores_gemma":[0.0054367995,0.0009320883,0.00128388,0.009004441,0.00048554837,0.002230007,0.0025388182,0.0019691864,0.17496826],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000039561834,0.00001745772,0.0004067224,0.000589715,0.000016285674,0.000021948068,0.000024119907,0.00017213746,0.00015863383,0.0004132329,0.99638075,0.0017594903],"study_design_scores_gemma":[0.000096980984,0.000012197764,0.0017467525,0.00018268201,0.000016201011,0.000055975564,0.000073936106,0.00028187956,0.0002816313,0.0009230791,0.99630827,0.000020322406],"about_ca_topic_score_codex":0.02125999,"about_ca_topic_score_gemma":0.036672287,"teacher_disagreement_score":0.88449365,"about_ca_system_score_codex":0.0017033621,"about_ca_system_score_gemma":0.0025607978,"threshold_uncertainty_score":0.38640732},"labels":[],"label_agreement":null},{"id":"W6925115210","doi":"10.1594/pangaea.325731","title":"Water temperature real-time profiles from cruise 2350899 (DCQC)","year":2005,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cruise; Temperature measurement; Sea surface temperature; Ectotherm","score_opus":0.016446779609527652,"score_gpt":0.2614480162291343,"score_spread":0.24500123661960663,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6925115210","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00020235169,0.00001710838,0.00003050541,0.000029714016,0.0000065152317,0.0000056172626,0.9992187,0.00017740215,0.00031203596],"genre_scores_gemma":[0.00042417768,0.00001630883,0.00012915855,0.000009811232,0.0000015013977,0.000020453555,0.9990715,0.000023842127,0.0003032832],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994134,0.00004864549,0.00006801115,0.0001543137,0.00021559797,0.00009999203],"domain_scores_gemma":[0.99849117,0.00019957757,0.00016949144,0.0003092884,0.00062318327,0.00020716361],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006290816,0.0018788985,0.0010444755,0.0029753791,0.0008240796,0.001222288,0.002796232,0.0017583006,0.024726104],"category_scores_gemma":[0.0032654868,0.00068737834,0.0011058094,0.007629683,0.0003595208,0.0009272475,0.0013187596,0.0013129133,0.03469225],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000080817816,0.00001995903,0.0017959866,0.00038083777,0.00003740014,0.00003223472,0.000035077665,0.00063043315,0.00021193044,0.0002146936,0.99442333,0.0021373702],"study_design_scores_gemma":[0.00033722457,0.000022436541,0.028363867,0.00026837547,0.000061250605,0.0000711891,0.0001972905,0.0013895062,0.0010580426,0.0006992081,0.9674676,0.00006396558],"about_ca_topic_score_codex":0.30484924,"about_ca_topic_score_gemma":0.41060117,"teacher_disagreement_score":0.30484924,"about_ca_system_score_codex":0.0028476624,"about_ca_system_score_gemma":0.004426171,"threshold_uncertainty_score":0.6061497},"labels":[],"label_agreement":null},{"id":"W6926490959","doi":"10.25384/sage.14964052","title":"sj-pdf-1-cjk-10.1177_20543581211027759 – Supplemental material for Incidence and Outcomes of Acute Kidney Injury in Patients Admitted to Hospital With COVID-19: A Retrospective Cohort Study","year":2021,"lang":"en","type":"article","venue":"Sage Journals Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Incidence (geometry); Retrospective cohort study; Acute kidney injury; Kidney disease; Cohort study; Cohort","score_opus":0.014729846460555708,"score_gpt":0.3317513048093872,"score_spread":0.3170214583488315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6926490959","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005563177,0.00013849618,0.0005262539,0.00057465106,0.00015008435,0.00018396665,0.9831535,0.0008761033,0.008833709],"genre_scores_gemma":[0.024693826,0.00034671123,0.0028488578,0.00077765953,0.00029302627,0.00071141205,0.9405578,0.0006781683,0.029092636],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99928564,0.00008327029,0.00019653588,0.00014825916,0.00018756483,0.00009868959],"domain_scores_gemma":[0.99084705,0.003695811,0.001336158,0.00085720274,0.0021050426,0.0011588312],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0008468965,0.00050589035,0.00065315695,0.003099501,0.00071699114,0.0016315449,0.0012957971,0.0013180841,0.575077],"category_scores_gemma":[0.009350637,0.0005596224,0.00057801296,0.0038338073,0.00017332437,0.0011680579,0.00097569387,0.00090513815,0.14980549],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006346645,0.0002387152,0.025140481,0.00062542583,0.00007140826,0.000266786,0.00006211458,0.00015858834,0.0007582691,0.00046037222,0.9490368,0.022546444],"study_design_scores_gemma":[0.00178839,0.00048413195,0.3124042,0.0013411812,0.00018277823,0.0027191953,0.0004354791,0.0015496891,0.0027939905,0.0037861064,0.6723139,0.0002009192],"about_ca_topic_score_codex":0.008638496,"about_ca_topic_score_gemma":0.011496548,"teacher_disagreement_score":0.575077,"about_ca_system_score_codex":0.0008104363,"about_ca_system_score_gemma":0.001628559,"threshold_uncertainty_score":0.60610104},"labels":[],"label_agreement":null},{"id":"W6926560065","doi":"10.25318/2710002901-fra","title":"Personnel de l'administration fédérale affecté aux activités scientifiques et technologiques, selon les principaux ministères et organismes - Perspectives","year":2019,"lang":"fr","type":"dataset","venue":"Statistics Canada Dissemination","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Work (physics); Context (archaeology); Derogation","score_opus":0.01454663586932905,"score_gpt":0.31345643564066317,"score_spread":0.29890979977133414,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6926560065","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004515353,0.000078115896,0.000057367877,0.000254654,0.000018072687,0.000010655537,0.99754024,0.00009328577,0.0014961795],"genre_scores_gemma":[0.0018176812,0.00010945108,0.00033393735,0.00006958076,0.00000937934,0.000053643078,0.99524915,0.000027400974,0.0023297702],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9977514,0.00022844154,0.00023548545,0.0004358052,0.00086406205,0.00048468655],"domain_scores_gemma":[0.99308395,0.0015431122,0.0006100988,0.00060024014,0.003369127,0.0007934645],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019126894,0.0013702511,0.0009875331,0.0059105977,0.0012241476,0.0024254646,0.0021236148,0.0011372973,0.014392687],"category_scores_gemma":[0.0121208485,0.00049186224,0.0009191424,0.012512375,0.0005343497,0.0009698531,0.001375777,0.002048043,0.011809236],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007468643,0.00001808122,0.0050639124,0.00030203935,0.000026107573,0.000015563994,0.000051431343,0.0002439592,0.0000604583,0.0009537322,0.98865753,0.004532553],"study_design_scores_gemma":[0.00015711221,0.000010860506,0.046133813,0.00030624738,0.000043612225,0.000039448147,0.0002753461,0.00059707795,0.00043651942,0.00065585686,0.95130956,0.000034583878],"about_ca_topic_score_codex":0.824213,"about_ca_topic_score_gemma":0.87964106,"teacher_disagreement_score":0.9847297,"about_ca_system_score_codex":0.015270267,"about_ca_system_score_gemma":0.023059506,"threshold_uncertainty_score":0.35364437},"labels":[],"label_agreement":null},{"id":"W6928879901","doi":"10.4000/14h9k","title":"2. L’avenir du pluralisme face aux majorités culturelles","year":2018,"lang":"fr","type":"book-chapter","venue":"Presses de l’Université de Montréal eBooks","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Face (sociological concept); Subject (documents); Context (archaeology); Politics","score_opus":0.012594015257825426,"score_gpt":0.20649943151729744,"score_spread":0.19390541625947202,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6928879901","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022287339,0.032914523,0.012536391,0.09161281,0.0037066194,0.000058477508,0.00028693632,0.00013778506,0.83645916],"genre_scores_gemma":[0.3000601,0.014641987,0.0069575333,0.016632775,0.0028115872,0.00014494522,0.00029350418,0.00040640548,0.6580512],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99747026,0.0008958236,0.00007803708,0.00050785363,0.00070466514,0.0003434008],"domain_scores_gemma":[0.99723744,0.0014873876,0.00017624784,0.0002840843,0.0005534606,0.00026130571],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0035277272,0.0007548842,0.00055267033,0.0017924741,0.0075459443,0.011486407,0.00081902783,0.0032603631,0.007096257],"category_scores_gemma":[0.004018028,0.00036115857,0.0005113539,0.0023718183,0.030517466,0.00821841,0.0034883306,0.006068249,0.0013919802],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007696569,0.0000052666746,0.00015676784,0.00007046639,0.0000035389717,0.000070671515,0.035604108,0.00005148454,0.00017537469,0.93416405,0.0185219,0.011168766],"study_design_scores_gemma":[0.0000040902987,0.000006087023,0.00076614163,0.00022673923,0.000004539403,0.00017989209,0.008630095,0.00011451813,0.00017384018,0.10391901,0.88596135,0.000013692679],"about_ca_topic_score_codex":0.144342,"about_ca_topic_score_gemma":0.17550506,"teacher_disagreement_score":0.144342,"about_ca_system_score_codex":0.016597426,"about_ca_system_score_gemma":0.01280388,"threshold_uncertainty_score":0.2870037},"labels":[],"label_agreement":null},{"id":"W6928975102","doi":"10.4231/d3rf5kg52","title":"Regional Seismic Hazard Assessment for Small Urban Centres in Western Canada","year":2014,"lang":"en","type":"article","venue":"Texas Advanced Computing Center","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Borehole; Bedrock; Seismic microzonation; Geological survey; Seismic hazard; Hazard; Seismic survey; Hazard analysis; Drilling","score_opus":0.01357739265115559,"score_gpt":0.2703017969604531,"score_spread":0.2567244043092975,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6928975102","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9781066,0.00052359526,0.00095579313,0.00028887924,0.0000109884495,0.00021641029,0.005569653,0.00011928572,0.014208827],"genre_scores_gemma":[0.9892887,0.00036643687,0.0010598315,0.00005025624,0.0000029243492,0.000038135815,0.0023185788,0.000014301802,0.0068608075],"study_design_codex":"observational","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9996619,0.000013781556,0.0000117044,0.00003090591,0.00015005555,0.0001316685],"domain_scores_gemma":[0.99934846,0.00002269821,0.00005832906,0.000014987471,0.00039895665,0.00015655195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00025007632,0.0005405943,0.00021679983,0.0023955978,0.0018533433,0.0010231512,0.0007410054,0.00017744151,0.0024729713],"category_scores_gemma":[0.00065943546,0.00021329471,0.00026906966,0.0036231468,0.00043209692,0.00017147603,0.00084482593,0.00022750883,0.0002650612],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040799254,0.00012283733,0.84536225,0.00022904627,0.0001492031,0.0011677023,0.0034010247,0.014672899,0.004445513,0.0012813715,0.013302964,0.115457214],"study_design_scores_gemma":[0.0000191559,0.0000576905,0.9758034,0.000056099183,0.00005221265,0.00013841172,0.0069738342,0.007317547,0.00067927974,0.00014967204,0.0087167155,0.000036014448],"about_ca_topic_score_codex":0.9955663,"about_ca_topic_score_gemma":0.99879503,"teacher_disagreement_score":0.024614312,"about_ca_system_score_codex":0.024614312,"about_ca_system_score_gemma":0.03055773,"threshold_uncertainty_score":0.17859018},"labels":[],"label_agreement":null},{"id":"W6929816825","doi":"10.5281/zenodo.11193046","title":"Iranattus principalis","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Taxonomy (biology); Carapace; Receptacle; Parakeet; Subgenus","score_opus":0.03423409405639634,"score_gpt":0.2756271101544633,"score_spread":0.24139301609806696,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6929816825","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.25472447,0.035738487,0.015416773,0.0015799176,0.0016126436,0.00092894514,0.006321799,0.0018385048,0.6818385],"genre_scores_gemma":[0.8327988,0.008453215,0.019387608,0.001612226,0.00058551045,0.0005226709,0.005593951,0.00019957635,0.13084643],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99969554,0.00003067149,0.000020441674,0.00011081291,0.00009120316,0.000051304178],"domain_scores_gemma":[0.9997694,0.000031232303,0.0000781079,0.000027215694,0.00006278188,0.000031305837],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00021365132,0.0011004641,0.0004752784,0.0026777314,0.0019002779,0.0005224016,0.00069870066,0.00051283167,0.02055418],"category_scores_gemma":[0.00052289525,0.00043448491,0.00033944042,0.0010780616,0.0011374191,0.001430989,0.0015889903,0.0013217403,0.009701864],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001005027,0.0002045439,0.02642225,0.0014015875,0.000121955774,0.0019086137,0.002496824,0.00088720483,0.045387983,0.025351543,0.05281858,0.84199387],"study_design_scores_gemma":[0.0001381478,0.00038363144,0.15699731,0.00094592187,0.00016551687,0.0062213195,0.0017563191,0.00070964376,0.004917725,0.0074704792,0.8202117,0.00008218723],"about_ca_topic_score_codex":0.010618157,"about_ca_topic_score_gemma":0.0156115275,"teacher_disagreement_score":0.02055418,"about_ca_system_score_codex":0.0009019532,"about_ca_system_score_gemma":0.00095110596,"threshold_uncertainty_score":0.068760574},"labels":[],"label_agreement":null},{"id":"W6929942944","doi":"10.5281/zenodo.10521711","title":"ome/ome-zarr-py: v0.4.1","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Tellabs (Canada)","funders":"","keywords":"Troubleshooting; Software; Work (physics); Table (database); Process (computing)","score_opus":0.027672940322213346,"score_gpt":0.26636422642681384,"score_spread":0.2386912861046005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6929942944","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0009780736,0.00024831324,0.077606834,0.00027623915,0.00019285263,0.00015491225,0.078141384,0.83000135,0.012400085],"genre_scores_gemma":[0.018649047,0.000421462,0.06452489,0.00065914576,0.00015169999,0.00078324333,0.3386341,0.55197346,0.024202965],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976426,0.00028437565,0.00022406018,0.00055002933,0.00083605055,0.00046283027],"domain_scores_gemma":[0.99734193,0.000605331,0.00019002108,0.00095730973,0.0006117913,0.00029366612],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003915657,0.0039180717,0.002424777,0.0028148927,0.0009949413,0.004915842,0.007108618,0.0030942939,0.22037394],"category_scores_gemma":[0.009699449,0.0027784954,0.002611676,0.0023803827,0.0011473693,0.0043966114,0.0060579493,0.003665573,0.3088584],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008883463,0.00009288404,0.0007760689,0.0011745334,0.00014445801,0.00018459129,0.00020759166,0.001759315,0.004974682,0.0071104486,0.9462809,0.03640624],"study_design_scores_gemma":[0.0006161557,0.0000728606,0.0014574504,0.00025517703,0.0000804437,0.00036208387,0.00009052813,0.01566119,0.021359464,0.0133575685,0.94644415,0.00024298081],"about_ca_topic_score_codex":0.006873963,"about_ca_topic_score_gemma":0.004535854,"teacher_disagreement_score":0.22037394,"about_ca_system_score_codex":0.0017313235,"about_ca_system_score_gemma":0.0016448573,"threshold_uncertainty_score":0.7372243},"labels":[],"label_agreement":null},{"id":"W6930240016","doi":"10.5281/zenodo.11457992","title":"~#@-&amp; [[[[WhatsApp ((( (+27) 736616875))) *____**)) TOP QUALITY COUNTERFEIT MONEY FOR SALE Barbados Killester Soweto\\\\EUROPE USA 3","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Counterfeit; Euros; Quality (philosophy); Counterfeit Drugs; Electronic money","score_opus":0.0483956539373101,"score_gpt":0.31079478049712655,"score_spread":0.26239912655981645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6930240016","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005330692,0.00014944571,0.00044108735,0.0010805513,0.0008304634,0.00006149516,0.0012912011,0.0018053614,0.99380726],"genre_scores_gemma":[0.0013274519,0.000099420125,0.00015943024,0.00020223048,0.00005138911,0.000010287677,0.00038565547,0.00039069392,0.99737334],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996427,0.000024528994,0.0000107446385,0.00006405709,0.00016955523,0.00008841068],"domain_scores_gemma":[0.9983803,0.00011335408,0.00006860215,0.00013376426,0.00074644515,0.0005575917],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00029684982,0.0006405117,0.00041739954,0.0008009635,0.0019723994,0.004639331,0.00060750335,0.0011927362,0.92407286],"category_scores_gemma":[0.002006373,0.00039479687,0.00034008047,0.00084687065,0.0005319519,0.0024456973,0.0019029232,0.0010769253,0.8896779],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011973477,0.00001237897,0.00012695444,0.00003037652,6.738379e-7,0.000026125903,0.000029759192,0.0000139672475,0.00020541366,0.0012174721,0.96877813,0.029546794],"study_design_scores_gemma":[0.0000028528,0.000007805674,0.0003671192,0.000018433797,6.0794963e-7,0.00003809513,0.000054465498,0.00002471966,0.00012873572,0.00012305482,0.99923027,0.0000037983448],"about_ca_topic_score_codex":0.006418729,"about_ca_topic_score_gemma":0.012781093,"teacher_disagreement_score":0.07592714,"about_ca_system_score_codex":0.0010031592,"about_ca_system_score_gemma":0.000980166,"threshold_uncertainty_score":0.108300745},"labels":[],"label_agreement":null},{"id":"W6930322954","doi":"10.5281/zenodo.12262476","title":"Burger king coupons pdf 2023 mdrz","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Meal; Coupon; Mail order; Living room; Order (exchange); Gable","score_opus":0.028845965310899016,"score_gpt":0.2675002027074825,"score_spread":0.23865423739658348,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6930322954","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00037037232,0.00033197005,0.0004167543,0.0007670119,0.0013609666,0.00012990713,0.0021117232,0.0017824719,0.9927288],"genre_scores_gemma":[0.00094397104,0.00013267854,0.00015792792,0.00021951606,0.000098370845,0.000020875583,0.0006382,0.0004413597,0.9973471],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992812,0.000038958846,0.00002318682,0.000071195755,0.00047376278,0.000111699104],"domain_scores_gemma":[0.9985682,0.00012263305,0.000039533716,0.00015708031,0.00084879494,0.0002638293],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0005349549,0.001051378,0.00078321557,0.0015766864,0.0024094698,0.005662848,0.001249522,0.0021968172,0.92494833],"category_scores_gemma":[0.0027964397,0.0008533038,0.00081636617,0.0010161725,0.00071699795,0.0033602333,0.0026043653,0.0028222264,0.85885],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027194532,0.000015642607,0.00004757284,0.00007104798,0.0000014196032,0.000052494273,0.000033730175,0.00004030864,0.0002835866,0.0019040573,0.97618365,0.021339154],"study_design_scores_gemma":[0.0000061609308,0.0000114154,0.00020955176,0.00003870286,0.0000011979351,0.00006242213,0.000056737215,0.00003315252,0.00021832931,0.00025227133,0.9991043,0.000005816887],"about_ca_topic_score_codex":0.007866853,"about_ca_topic_score_gemma":0.015647871,"teacher_disagreement_score":0.075051665,"about_ca_system_score_codex":0.0019179463,"about_ca_system_score_gemma":0.0011746443,"threshold_uncertainty_score":0.10705203},"labels":[],"label_agreement":null},{"id":"W6930686379","doi":"10.5281/zenodo.1288768","title":"Outcome Of Patients With Acute Left Ventricular Faliure After Coronary Artery Bypass Grafting","year":2018,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Artery; Bypass grafting; Coronary artery bypass surgery; Acute kidney injury; Observational study; Stroke (engine)","score_opus":0.014731524303851332,"score_gpt":0.23799761380230022,"score_spread":0.2232660894984489,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6930686379","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9994837,0.00010551215,0.000018900155,0.00006760496,0.000008438882,0.0000051874454,0.000057596524,8.332469e-7,0.00025228556],"genre_scores_gemma":[0.99969804,0.000060497267,0.000015458503,0.000025262825,0.00001686091,0.0000045826782,0.00012427957,3.620717e-7,0.00005474106],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99972135,0.00004916069,0.000040271727,0.0000390203,0.000047806756,0.00010237744],"domain_scores_gemma":[0.9992531,0.00009279545,0.00034200444,0.000022665114,0.000058053352,0.00023144514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0003186645,0.00016198463,0.00028569112,0.00043495005,0.00050150475,0.00046257532,0.00021927492,0.00036416037,0.0015258379],"category_scores_gemma":[0.0014821459,0.00013123549,0.00028464032,0.00036052044,0.0002730565,0.00038797624,0.0004609799,0.0005407293,0.00013531018],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00012771187,0.00007741309,0.99775827,0.000013611586,0.000017445516,0.0006487076,0.00008815031,0.000023449284,0.00009385297,0.00002049305,0.00012913499,0.0010017118],"study_design_scores_gemma":[0.000010634241,0.0003954587,0.9971316,0.000020146861,0.000017794096,0.0014640132,0.0006004151,0.0001146817,0.00005579282,0.00005283416,0.00013003346,0.0000065893755],"about_ca_topic_score_codex":0.0011391975,"about_ca_topic_score_gemma":0.0021233128,"teacher_disagreement_score":0.0015258379,"about_ca_system_score_codex":0.000339674,"about_ca_system_score_gemma":0.0003882276,"threshold_uncertainty_score":0.005104482},"labels":[],"label_agreement":null},{"id":"W6931183612","doi":"10.5281/zenodo.5695254","title":"Craspedolepta fumida Caldwell 1938","year":2012,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Natural (archaeology); Derogation; Race (biology); Context (archaeology)","score_opus":0.03587282349802403,"score_gpt":0.2642340053184256,"score_spread":0.22836118182040158,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931183612","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.586509,0.014989607,0.010857294,0.0010480229,0.0005191077,0.00046880366,0.011052091,0.0008865632,0.37366945],"genre_scores_gemma":[0.94184613,0.003830803,0.0057489397,0.0005905363,0.000071595714,0.000101625745,0.0036148038,0.00003699986,0.044158477],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9999224,0.0000047267695,0.0000042854003,0.000033520264,0.000022222517,0.000012907558],"domain_scores_gemma":[0.9998184,0.000040564704,0.000046804842,0.000013679865,0.000057410234,0.000023085622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000068713256,0.0004674735,0.00018357877,0.0012993892,0.0017626558,0.000335554,0.00046845377,0.00040684643,0.006838092],"category_scores_gemma":[0.00026658791,0.00018050274,0.000096046344,0.0007173414,0.00034910376,0.00046316948,0.0004757736,0.00040569965,0.0019761738],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00044854693,0.00012567513,0.077623665,0.0008656913,0.000075636446,0.0041863415,0.002174943,0.0007955611,0.12098862,0.005367066,0.035966568,0.75138164],"study_design_scores_gemma":[0.0000787183,0.00018070838,0.35801008,0.00038708834,0.00007609772,0.0032550627,0.0016285715,0.0006069594,0.011638035,0.0010917115,0.6230125,0.000034427474],"about_ca_topic_score_codex":0.08111369,"about_ca_topic_score_gemma":0.22359842,"teacher_disagreement_score":0.08111369,"about_ca_system_score_codex":0.0012236324,"about_ca_system_score_gemma":0.0005451974,"threshold_uncertainty_score":0.16128314},"labels":[],"label_agreement":null},{"id":"W6931266175","doi":"10.5281/zenodo.5746043","title":"TraceSim: An Alignment Method for Computing Stack Trace Similarity","year":2022,"lang":"en","type":"article","venue":"PolyPublie (École Polytechnique de Montréal)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Polytechnique Montréal","funders":"","keywords":"TRACE (psycholinguistics); Stack (abstract data type); Similarity (geometry); Similitude; Pattern recognition (psychology); George (robot)","score_opus":0.020327516886866636,"score_gpt":0.29791859765804724,"score_spread":0.2775910807711806,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931266175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0065324223,0.0003880673,0.7286774,0.00025654552,0.00036319933,0.00040595385,0.043022066,0.21500643,0.0053479364],"genre_scores_gemma":[0.062139317,0.0005846237,0.67547196,0.00019097209,0.00018891935,0.0014441122,0.20092997,0.050592326,0.008457744],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9954124,0.0006104161,0.0008226849,0.00089781254,0.0019506741,0.00030601496],"domain_scores_gemma":[0.99359506,0.0020172305,0.00054284715,0.0016433316,0.001925706,0.00027581913],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027800947,0.0028853759,0.0015878364,0.010368176,0.0019222777,0.0037943725,0.0037235885,0.0018940701,0.032839213],"category_scores_gemma":[0.026436593,0.0014923747,0.0023671743,0.011299448,0.00079788436,0.0065924446,0.0047691935,0.0023711415,0.02308336],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013438498,0.00033873424,0.008063418,0.0022602137,0.0005309441,0.00061392965,0.0013685062,0.015060196,0.0185825,0.036727127,0.37427583,0.5408348],"study_design_scores_gemma":[0.00048291023,0.00044080173,0.0068090577,0.0005334671,0.00039096968,0.0011220389,0.0012497162,0.27003965,0.060238674,0.12725854,0.53103447,0.00039971375],"about_ca_topic_score_codex":0.006056081,"about_ca_topic_score_gemma":0.008991104,"teacher_disagreement_score":0.032839213,"about_ca_system_score_codex":0.0011941077,"about_ca_system_score_gemma":0.0040271766,"threshold_uncertainty_score":0.109858096},"labels":[],"label_agreement":null},{"id":"W6931275717","doi":"10.5281/zenodo.4609334","title":"CINECA_Query expansion service_D1.2","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"European Commission","keywords":"Discoverability; Ontology; Representation (politics); SPARQL; Ranking (information retrieval); External Data Representation; Data access; Data integration; RDF","score_opus":0.04041229772472648,"score_gpt":0.2513374146380585,"score_spread":0.210925116913332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931275717","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033730287,0.00049677474,0.17106208,0.003612976,0.0005946871,0.0012195127,0.08486801,0.6825712,0.052201763],"genre_scores_gemma":[0.10518052,0.0014290293,0.25979754,0.014338276,0.0008163243,0.0034854857,0.39189553,0.15372065,0.06933663],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99340516,0.0013463778,0.00069827883,0.0014037655,0.002573255,0.0005732625],"domain_scores_gemma":[0.98564756,0.004785909,0.00037646346,0.005327173,0.0031258503,0.00073706947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008846716,0.002734361,0.0018215029,0.0032779905,0.0016162059,0.0074059544,0.004547221,0.0030415852,0.07916607],"category_scores_gemma":[0.024321146,0.001711697,0.0031265526,0.0029796087,0.0015626015,0.00930788,0.009814223,0.004266407,0.052178543],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010416937,0.00016762903,0.0030512859,0.00062307995,0.00013626272,0.00042936607,0.000833954,0.0015292775,0.00592775,0.025059037,0.90123975,0.059960935],"study_design_scores_gemma":[0.0003653976,0.000089381996,0.002691624,0.00017622994,0.000044982775,0.00048067406,0.00039395972,0.038871706,0.013253939,0.021196805,0.922245,0.00019028704],"about_ca_topic_score_codex":0.033549044,"about_ca_topic_score_gemma":0.017307887,"teacher_disagreement_score":0.07916607,"about_ca_system_score_codex":0.0035089734,"about_ca_system_score_gemma":0.0045408043,"threshold_uncertainty_score":0.2648369},"labels":[],"label_agreement":null},{"id":"W6931359040","doi":"10.5281/zenodo.7750938","title":"Supplementary material: Inferring clusters of orthologous and paralogous transcripts","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Sherbrooke","funders":"","keywords":"Genome; Sequence (biology); Human genome; Genomics","score_opus":0.03270626207176303,"score_gpt":0.2641647382390424,"score_spread":0.23145847616727938,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931359040","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00053272594,0.00008030512,0.00964609,0.00014122942,0.00013010787,0.000056058,0.9780001,0.008918874,0.002494566],"genre_scores_gemma":[0.0041570817,0.00015153155,0.0265605,0.00014353398,0.000084473,0.00025858375,0.9609492,0.004299292,0.003395765],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999141,0.000121272955,0.00013127472,0.00027701264,0.00025404888,0.00007544399],"domain_scores_gemma":[0.99233925,0.005192092,0.00033069297,0.00093423884,0.0009349893,0.0002687169],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0013753355,0.002000105,0.0011235949,0.004542447,0.0007869131,0.002428226,0.0017684081,0.0016022332,0.5858246],"category_scores_gemma":[0.012784596,0.00094596954,0.0012979012,0.004917984,0.00039216457,0.0020657463,0.0014126631,0.001131197,0.19932903],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026879975,0.000074099844,0.0015247884,0.0022721167,0.00009352664,0.00021192507,0.00008745729,0.0016180164,0.0030582487,0.002680001,0.96007454,0.028036483],"study_design_scores_gemma":[0.0008253281,0.000107530934,0.009456277,0.0007442497,0.00019400365,0.0012556274,0.00026333993,0.012898478,0.009627752,0.025785642,0.9387187,0.0001231639],"about_ca_topic_score_codex":0.0033949923,"about_ca_topic_score_gemma":0.005343119,"teacher_disagreement_score":0.5858246,"about_ca_system_score_codex":0.00092700595,"about_ca_system_score_gemma":0.0015263981,"threshold_uncertainty_score":0.5907709},"labels":[],"label_agreement":null},{"id":"W6931377624","doi":"10.5281/zenodo.6273272","title":"Epigamia magna Berkeley 1923, comb. n.","year":2004,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Dorsum; Holotype; GenBank; Orange (colour); Sequence (biology)","score_opus":0.025047505014001478,"score_gpt":0.25105123195362355,"score_spread":0.22600372693962206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931377624","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09844334,0.0620364,0.0042188056,0.0011264792,0.0014210382,0.00071077526,0.011183376,0.0012336218,0.81962615],"genre_scores_gemma":[0.7879097,0.023464818,0.0085928105,0.0018616393,0.001415927,0.00051098206,0.008609249,0.00026833502,0.1673666],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9998385,0.000013426409,0.000009749605,0.000056325873,0.000056961784,0.000025033958],"domain_scores_gemma":[0.9998342,0.000032094908,0.00004758235,0.00001532326,0.000049308706,0.000021496804],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00014898984,0.000762928,0.0003905413,0.0028489355,0.0016767906,0.000471995,0.0006107201,0.00035479642,0.02439369],"category_scores_gemma":[0.0003694142,0.00038534234,0.00020708641,0.0013489126,0.0009121964,0.0008768506,0.0008870491,0.00065516506,0.006633677],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00033792146,0.000056598234,0.01888836,0.0004825591,0.00005320392,0.0005366146,0.0015479437,0.00026128438,0.005749951,0.0063717514,0.10613858,0.8595753],"study_design_scores_gemma":[0.00006845786,0.00006566153,0.27233034,0.00023629013,0.00006381108,0.0023362536,0.0005955626,0.0001567928,0.00047271838,0.0012037128,0.72243893,0.000031494896],"about_ca_topic_score_codex":0.033257272,"about_ca_topic_score_gemma":0.06481576,"teacher_disagreement_score":0.033257272,"about_ca_system_score_codex":0.0011194588,"about_ca_system_score_gemma":0.0005344939,"threshold_uncertainty_score":0.08160496},"labels":[],"label_agreement":null},{"id":"W6931670828","doi":"10.5683/sp3/tezaq9","title":"Recensements du Canada 1665-1871, Canada atlantique - Acadie","year":2023,"lang":"fr","type":"dataset","venue":"Borealis","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Statistical analysis; Research methodology","score_opus":0.013128012277344897,"score_gpt":0.2436863510999396,"score_spread":0.23055833882259472,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931670828","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08453895,0.025686143,0.001851741,0.008550894,0.0011896034,0.00020371327,0.6096183,0.0005189348,0.26784167],"genre_scores_gemma":[0.34817863,0.029392567,0.0048942477,0.0017063943,0.00032882643,0.00025528256,0.20887886,0.0005166596,0.40584847],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99791723,0.00006450949,0.00010283442,0.00025777164,0.0013006402,0.00035705577],"domain_scores_gemma":[0.99556446,0.00037442966,0.00024438885,0.00012782445,0.003374794,0.0003141425],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001019027,0.00051663903,0.0004935493,0.012354548,0.0052509103,0.0040793,0.0008820873,0.00052215985,0.032158453],"category_scores_gemma":[0.0046504154,0.0003468868,0.00045362595,0.034288008,0.0014076185,0.0012168944,0.0014743704,0.0011133872,0.004007336],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00037321527,0.000047957517,0.08671597,0.0019105624,0.00019590779,0.0012978851,0.028246442,0.0014426778,0.0011297825,0.035351932,0.6061145,0.2371732],"study_design_scores_gemma":[0.0000071768645,0.000007425006,0.17905538,0.00037896758,0.000029231844,0.000082094615,0.0063976496,0.00009864704,0.00057376496,0.000449631,0.8128784,0.000041699495],"about_ca_topic_score_codex":0.99543226,"about_ca_topic_score_gemma":0.99763966,"teacher_disagreement_score":0.065564305,"about_ca_system_score_codex":0.065564305,"about_ca_system_score_gemma":0.09578682,"threshold_uncertainty_score":0.47570455},"labels":[],"label_agreement":null},{"id":"W6931868449","doi":"10.5281/zenodo.7629415","title":"Fig. 30 in Revision of the Nearctic species of the Lasioglossum (Dialictus) gemmatum species complex (Hymenoptera: Halictidae)","year":2023,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Holotype; Nearctic ecozone; Species complex; Taxonomy (biology); Scale (ratio)","score_opus":0.044680809145914616,"score_gpt":0.25958469399421663,"score_spread":0.21490388484830203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6931868449","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.026464535,0.014416191,0.019803641,0.004793164,0.015118676,0.00077634177,0.14295423,0.0055387225,0.77013445],"genre_scores_gemma":[0.10883712,0.0134818135,0.056999892,0.0017045329,0.0027386004,0.0009301106,0.22306007,0.0038411927,0.5884067],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997334,0.000045375185,0.000034961384,0.000066727414,0.00008958704,0.000029929906],"domain_scores_gemma":[0.9995634,0.00006280756,0.00005115694,0.00005910756,0.00021260254,0.0000509316],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006279432,0.00074048055,0.00031937344,0.004417105,0.0009672872,0.0009121308,0.0007107445,0.000444489,0.08133543],"category_scores_gemma":[0.0011829237,0.00021718383,0.00031485735,0.0029343965,0.0008010997,0.0011131184,0.00072671776,0.0009333164,0.035803355],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018460465,0.000029277277,0.003311099,0.000969766,0.000025666031,0.0002541125,0.00078605296,0.00024043187,0.0036580667,0.008057147,0.7475959,0.23488778],"study_design_scores_gemma":[0.0000059774857,0.000008910352,0.006090645,0.000108657,0.000012495714,0.00023788953,0.00013640955,0.000055899247,0.000258729,0.00045321698,0.992626,0.0000050289655],"about_ca_topic_score_codex":0.015121472,"about_ca_topic_score_gemma":0.031585246,"teacher_disagreement_score":0.08133543,"about_ca_system_score_codex":0.0011781887,"about_ca_system_score_gemma":0.0010277457,"threshold_uncertainty_score":0.27209413},"labels":[],"label_agreement":null},{"id":"W6939193458","doi":"10.60692/zzk1n-dwn51","title":"Developing data interoperability using standards: A wheat community use case","year":2017,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Institut National de la Recherche Scientifique","funders":"","keywords":"Interoperability; Semantic interoperability; Metadata; Ontology; Data exchange; Cross-domain interoperability; Linked data; The Internet","score_opus":0.25977750430107655,"score_gpt":0.3475690416771639,"score_spread":0.08779153737608736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939193458","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.32850078,0.002670652,0.4733768,0.07155151,0.00048117008,0.0024225817,0.0006161435,0.0015917822,0.11878862],"genre_scores_gemma":[0.5491925,0.0022921928,0.42561385,0.0056220596,0.00017790812,0.0010505648,0.0018127601,0.0007817424,0.013456447],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9055215,0.060009424,0.0066290437,0.0040088794,0.020120548,0.0037106303],"domain_scores_gemma":[0.885159,0.06466247,0.0038715077,0.018248875,0.024472766,0.0035853875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.11503063,0.0010130208,0.0009096538,0.0054762354,0.009394334,0.015989417,0.0054030577,0.010764041,0.0025447637],"category_scores_gemma":[0.07486249,0.0013357805,0.0021118643,0.008594692,0.007371901,0.034058988,0.018603228,0.0067604496,0.00080944377],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00029363154,0.0019997167,0.035354186,0.0014887478,0.00017252451,0.014348539,0.11205558,0.0075035384,0.009099288,0.51919466,0.028952334,0.26953727],"study_design_scores_gemma":[0.00022810744,0.00084500434,0.008200962,0.0025370943,0.0001882214,0.006300985,0.09732672,0.03898334,0.019799333,0.14041045,0.6848069,0.00037292764],"about_ca_topic_score_codex":0.017378792,"about_ca_topic_score_gemma":0.017321907,"teacher_disagreement_score":0.11503063,"about_ca_system_score_codex":0.008208195,"about_ca_system_score_gemma":0.00940439,"threshold_uncertainty_score":0.6083474},"labels":[],"label_agreement":null},{"id":"W6939297872","doi":"10.60692/c3asp-twh58","title":"Future-proofing and maximizing the utility of metadata: The PHA4GE SARS-CoV-2 contextual data specification package","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; Dalhousie University; McMaster University; BC Centre for Disease Control; Simon Fraser University","funders":"","keywords":"Interoperability; Consistency (knowledge bases); Standardization; Harmonization; Contextual design; Data integration; Data aggregator; Data quality; Data consistency; Alliance","score_opus":0.10648704050317989,"score_gpt":0.27160296818766494,"score_spread":0.16511592768448505,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939297872","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0053257,0.0004026599,0.8932614,0.0063521247,0.0005850911,0.0018623208,0.034880105,0.04214803,0.01518253],"genre_scores_gemma":[0.025218792,0.00080293027,0.8606451,0.0027043456,0.0002642512,0.0020355026,0.090431765,0.010413766,0.0074835303],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.98052067,0.007575734,0.0039986656,0.001453796,0.0054073473,0.0010438622],"domain_scores_gemma":[0.95510644,0.012810058,0.0027479504,0.017408175,0.010161257,0.0017661541],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.040992353,0.0014732435,0.0009528674,0.0044257403,0.0017106859,0.006855598,0.0038054178,0.0024595452,0.009659569],"category_scores_gemma":[0.0544225,0.0016045116,0.002537274,0.0034584226,0.0019048713,0.007890446,0.008617221,0.0039484813,0.009558591],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008619076,0.0004322967,0.013782091,0.0018657506,0.00024255374,0.0010296183,0.0032097301,0.013029661,0.015808472,0.3257431,0.38810214,0.23589277],"study_design_scores_gemma":[0.00014738046,0.00012993076,0.0023974339,0.0012110089,0.00009336534,0.000521788,0.00070055,0.016972678,0.015130654,0.08299392,0.87947524,0.00022614063],"about_ca_topic_score_codex":0.013373684,"about_ca_topic_score_gemma":0.012236432,"teacher_disagreement_score":0.9590076,"about_ca_system_score_codex":0.002570522,"about_ca_system_score_gemma":0.012620243,"threshold_uncertainty_score":0.21679085},"labels":[],"label_agreement":null},{"id":"W6939404866","doi":"10.6084/m9.figshare.20175968.v1","title":"Additional file 14 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07399661812176367,"score_gpt":0.34046377603794686,"score_spread":0.2664671579161832,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939404866","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000114003065,0.0000164146,0.0009859488,0.00012774351,0.000049567185,0.000051374034,0.9941298,0.0019005141,0.0026246323],"genre_scores_gemma":[0.0039529074,0.00016624697,0.009945814,0.00037870576,0.00007804421,0.0006671069,0.9700561,0.004763069,0.00999198],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994462,0.00008333833,0.00009016223,0.0001364405,0.0001518501,0.00009185429],"domain_scores_gemma":[0.9894957,0.0073070847,0.0004817509,0.0007402499,0.0016168051,0.0003583944],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001625849,0.0011890245,0.0010587021,0.0037770495,0.00085339893,0.0025318675,0.0019275527,0.0012354072,0.8335072],"category_scores_gemma":[0.016101614,0.0007714233,0.001374669,0.0054687276,0.00046558035,0.0034788605,0.0017997168,0.001367432,0.29080993],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014386566,0.000042544052,0.00055869337,0.0011516939,0.00002639132,0.00004790204,0.00007566005,0.00025463747,0.00017946526,0.0018051475,0.9884906,0.007223458],"study_design_scores_gemma":[0.0010048961,0.000033171204,0.003677951,0.00078367937,0.000064710635,0.00019373735,0.00022878159,0.0010887467,0.0011072819,0.01199134,0.9797489,0.00007674655],"about_ca_topic_score_codex":0.00976216,"about_ca_topic_score_gemma":0.012728693,"teacher_disagreement_score":0.8335072,"about_ca_system_score_codex":0.0016808321,"about_ca_system_score_gemma":0.0023606431,"threshold_uncertainty_score":0.23748171},"labels":[],"label_agreement":null},{"id":"W6939410549","doi":"10.6084/m9.figshare.14684514.v1","title":"Additional file 1 of Text mining to support abstract screening for knowledge syntheses: a semi-automated workflow","year":2021,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Queen's University; Toronto Metropolitan University","funders":"","keywords":"Workflow; Knowledge extraction; Text mining; Key (lock); File format; Web mining","score_opus":0.06223079883875487,"score_gpt":0.3101408277554949,"score_spread":0.24791002891674002,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939410549","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002243473,0.000033309032,0.002533689,0.0001704989,0.000054430577,0.00015518225,0.9908591,0.004225617,0.0017438536],"genre_scores_gemma":[0.0045948513,0.00015298907,0.023786917,0.00055599137,0.00014457376,0.0016701636,0.9542812,0.005170115,0.0096430965],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9987955,0.00021264068,0.00026578616,0.00035693744,0.00027168484,0.00009746074],"domain_scores_gemma":[0.9730911,0.020470308,0.0012916324,0.0015015325,0.0029710562,0.0006743365],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0026597513,0.0017180206,0.0011270446,0.0043257703,0.0009108515,0.0027135604,0.0021404906,0.0012907898,0.7997494],"category_scores_gemma":[0.029001778,0.00079491216,0.0012257409,0.0046520424,0.0004983741,0.002598365,0.0025680133,0.0012877261,0.2792254],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00038606586,0.00006791565,0.0011147693,0.0031474023,0.000055983113,0.00012865235,0.000097619624,0.00036761802,0.00061552017,0.00093373115,0.97473776,0.018346962],"study_design_scores_gemma":[0.0011053892,0.00011856013,0.006480332,0.0012312124,0.00015182115,0.00049852073,0.00027310092,0.0020981962,0.0033189899,0.012979848,0.9716176,0.00012651678],"about_ca_topic_score_codex":0.0029059334,"about_ca_topic_score_gemma":0.004861725,"teacher_disagreement_score":0.99734026,"about_ca_system_score_codex":0.0012082447,"about_ca_system_score_gemma":0.0025725313,"threshold_uncertainty_score":0.28563315},"labels":[],"label_agreement":null},{"id":"W6939546117","doi":"10.6084/m9.figshare.20175962.v1","title":"Additional file 12 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.08169920362507627,"score_gpt":0.33932481622852606,"score_spread":0.2576256126034498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939546117","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000112309484,0.000016735188,0.0009400864,0.00013280449,0.000052637144,0.000052426927,0.9944635,0.0017929341,0.0024364458],"genre_scores_gemma":[0.003789681,0.00016380964,0.00961402,0.00040781221,0.00008227296,0.00068683806,0.97090864,0.0045697438,0.009777189],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994442,0.00008451382,0.00009179797,0.0001389801,0.00014698821,0.000093549796],"domain_scores_gemma":[0.98917997,0.007571394,0.0004898853,0.00076281937,0.0016077715,0.00038827394],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016481322,0.0012226275,0.0011015199,0.0037755694,0.00088324153,0.0025400151,0.0019453253,0.0013122045,0.8303303],"category_scores_gemma":[0.016585866,0.00077525765,0.0014243278,0.005323235,0.00046743295,0.003507116,0.0018179094,0.0014214335,0.28613603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001623206,0.000044385546,0.0005657286,0.0012939303,0.000028011294,0.000054244894,0.00007722742,0.00025914604,0.00018882353,0.0017715802,0.988325,0.0072296634],"study_design_scores_gemma":[0.0010542165,0.000035899444,0.0037384492,0.00083275745,0.00006861288,0.00020843818,0.0002323219,0.001046053,0.0011038248,0.011816004,0.97978354,0.00007992351],"about_ca_topic_score_codex":0.0092043895,"about_ca_topic_score_gemma":0.012233973,"teacher_disagreement_score":0.8303303,"about_ca_system_score_codex":0.0016311294,"about_ca_system_score_gemma":0.002342848,"threshold_uncertainty_score":0.24201322},"labels":[],"label_agreement":null},{"id":"W6939627889","doi":"10.6084/m9.figshare.20175962","title":"Additional file 12 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.08169920362507627,"score_gpt":0.33932481622852606,"score_spread":0.2576256126034498,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939627889","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000112309484,0.000016735188,0.0009400864,0.00013280449,0.000052637144,0.000052426927,0.9944635,0.0017929341,0.0024364458],"genre_scores_gemma":[0.003789681,0.00016380964,0.00961402,0.00040781221,0.00008227296,0.00068683806,0.97090864,0.0045697438,0.009777189],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994442,0.00008451382,0.00009179797,0.0001389801,0.00014698821,0.000093549796],"domain_scores_gemma":[0.98917997,0.007571394,0.0004898853,0.00076281937,0.0016077715,0.00038827394],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016481322,0.0012226275,0.0011015199,0.0037755694,0.00088324153,0.0025400151,0.0019453253,0.0013122045,0.8303303],"category_scores_gemma":[0.016585866,0.00077525765,0.0014243278,0.005323235,0.00046743295,0.003507116,0.0018179094,0.0014214335,0.28613603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001623206,0.000044385546,0.0005657286,0.0012939303,0.000028011294,0.000054244894,0.00007722742,0.00025914604,0.00018882353,0.0017715802,0.988325,0.0072296634],"study_design_scores_gemma":[0.0010542165,0.000035899444,0.0037384492,0.00083275745,0.00006861288,0.00020843818,0.0002323219,0.001046053,0.0011038248,0.011816004,0.97978354,0.00007992351],"about_ca_topic_score_codex":0.0092043895,"about_ca_topic_score_gemma":0.012233973,"teacher_disagreement_score":0.8303303,"about_ca_system_score_codex":0.0016311294,"about_ca_system_score_gemma":0.002342848,"threshold_uncertainty_score":0.24201322},"labels":[],"label_agreement":null},{"id":"W6939663681","doi":"10.6084/m9.figshare.20175965","title":"Additional file 13 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07371233409508435,"score_gpt":0.33956229448162595,"score_spread":0.26584996038654163,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6939663681","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000109568005,0.000015301443,0.0009218042,0.00012244754,0.000048345446,0.000050967268,0.9944992,0.0018042241,0.0024282115],"genre_scores_gemma":[0.0036990028,0.00015464306,0.009229815,0.00038031,0.00007633109,0.0006796872,0.9717325,0.004598323,0.009449419],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99941754,0.00008769035,0.00009691767,0.00014391664,0.00015722231,0.00009675495],"domain_scores_gemma":[0.98882335,0.007814384,0.00050438126,0.0007720629,0.0016996528,0.00038606883],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001682783,0.0012212453,0.0010938367,0.0037570088,0.0008876773,0.0025859997,0.0019603868,0.0012781267,0.83599436],"category_scores_gemma":[0.016874032,0.000782945,0.0014015958,0.0054499856,0.0004749772,0.0035682686,0.001846981,0.0014088403,0.2936053],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015759983,0.000044137687,0.0005732425,0.0012435492,0.000027740585,0.000051372874,0.00007788562,0.00025874472,0.00018593013,0.0017527008,0.9885173,0.0071096714],"study_design_scores_gemma":[0.0010946946,0.000036381112,0.0038308438,0.000830316,0.00006953991,0.0002015288,0.00024099741,0.0010933084,0.0011698154,0.012214053,0.9791374,0.000081038896],"about_ca_topic_score_codex":0.009747807,"about_ca_topic_score_gemma":0.0129089225,"teacher_disagreement_score":0.83599436,"about_ca_system_score_codex":0.0016812171,"about_ca_system_score_gemma":0.002414099,"threshold_uncertainty_score":0.23393404},"labels":[],"label_agreement":null},{"id":"W6943354257","doi":"10.15468/dl.yz5asq","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6943354257","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008844714,0.000046266614,0.00007602837,0.000051301486,0.00001467343,0.000008461013,0.99802375,0.0007770696,0.00091394933],"genre_scores_gemma":[0.00017981396,0.00003748275,0.0002758469,0.00004398981,0.0000027844708,0.000035132794,0.9988456,0.00014256273,0.00043664651],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989417,0.00014059484,0.00014385402,0.00036757014,0.00025211438,0.00015419057],"domain_scores_gemma":[0.997969,0.00054528186,0.00019160808,0.0005286918,0.0005064074,0.00025897144],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00092235237,0.002456664,0.0014820251,0.004770391,0.0010435957,0.0023732448,0.0030488216,0.002047708,0.10018056],"category_scores_gemma":[0.005276086,0.0008823171,0.0012831356,0.008762912,0.00047253162,0.0021992382,0.0024468831,0.0019398679,0.15456133],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038878177,0.000017551682,0.00043777053,0.00051712774,0.000016132646,0.000022793463,0.000025692718,0.00017449349,0.00015448508,0.0004707736,0.9962901,0.0018342191],"study_design_scores_gemma":[0.00008320014,0.000011046772,0.0017859542,0.00016161069,0.000015002833,0.00006150176,0.000074120995,0.00029688125,0.0002809229,0.00090780784,0.9963026,0.000019393958],"about_ca_topic_score_codex":0.02528589,"about_ca_topic_score_gemma":0.04439353,"teacher_disagreement_score":0.89981943,"about_ca_system_score_codex":0.0017665672,"about_ca_system_score_gemma":0.0025830409,"threshold_uncertainty_score":0.33513737},"labels":[],"label_agreement":null},{"id":"W6945177364","doi":"10.25316/ir-14216","title":"Nanaimo Free Press [Saturday, February 21, 1880]","year":2020,"lang":"en","type":"other","venue":"VIURRSpace (Vancouver Island University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.010300429067392793,"score_gpt":0.20851102722341142,"score_spread":0.19821059815601863,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6945177364","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00047928133,0.0031791497,0.00022735648,0.0016199836,0.0018149291,0.000023612412,0.0050801816,0.000527472,0.987048],"genre_scores_gemma":[0.0005852764,0.0004471565,0.000067919405,0.00009190872,0.000052265823,0.0000061362734,0.00061614125,0.0001336393,0.9979996],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997266,0.00001700841,0.000009672208,0.0000485934,0.0001424947,0.000055608565],"domain_scores_gemma":[0.99970263,0.0000314996,0.000013508756,0.00002199144,0.00014207809,0.00008827982],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00031554396,0.0008583111,0.0005696871,0.0017882184,0.0024474137,0.0053489385,0.00065848185,0.001309214,0.470461],"category_scores_gemma":[0.0013295196,0.00037828993,0.00037019508,0.0027098944,0.00046100692,0.0021804257,0.0013665278,0.0015659163,0.28166845],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002337447,0.000006090388,0.00011172732,0.000051799605,0.000002214508,0.00004480318,0.000041522046,0.00004576918,0.00009042537,0.005185052,0.9577458,0.036651365],"study_design_scores_gemma":[0.0000017933335,0.000001696426,0.0002963375,0.000030986514,6.1295924e-7,0.000012765687,0.000017815604,0.000019583566,0.00003348451,0.00030085488,0.9992817,0.0000023523025],"about_ca_topic_score_codex":0.07206441,"about_ca_topic_score_gemma":0.29723755,"teacher_disagreement_score":0.9279356,"about_ca_system_score_codex":0.0022841133,"about_ca_system_score_gemma":0.002022191,"threshold_uncertainty_score":0.7553231},"labels":[],"label_agreement":null},{"id":"W6945301746","doi":"10.24411/2078-1318-2018-12134","title":"Характеристика высокопродуктивных коров в \"СХПК им. Кирова\" Кировской области","year":2018,"lang":"ru","type":"article","venue":"CyberLeninK (CyberLeninka)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Dairy cattle; Agriculture; Herd; Animal breeding; Dairy farming","score_opus":0.015359106994237775,"score_gpt":0.2774526002192771,"score_spread":0.26209349322503933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6945301746","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04266694,0.010996161,0.029529598,0.013868571,0.0015000822,0.00020035573,0.0011528247,0.00040927719,0.89967614],"genre_scores_gemma":[0.61282223,0.014137294,0.03445144,0.002393072,0.00075408636,0.0004534538,0.0015159026,0.00047721015,0.33299536],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.99614406,0.0008054001,0.00020387022,0.00066919083,0.0017029584,0.00047448056],"domain_scores_gemma":[0.9973538,0.0006331336,0.00031460452,0.00032536068,0.0010037774,0.00036936926],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0020774058,0.00064735813,0.00041400065,0.0022355618,0.003608695,0.009537028,0.00095291506,0.0019286639,0.045718446],"category_scores_gemma":[0.0045992746,0.00051787583,0.0007257409,0.0024253111,0.0038166125,0.0044325506,0.0029421567,0.0022638761,0.013274812],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015187719,0.0001100552,0.008829474,0.00073403,0.000059451013,0.0008053094,0.010268684,0.00095348223,0.004050949,0.7227331,0.046766754,0.20453693],"study_design_scores_gemma":[0.000021302187,0.00005804451,0.011388422,0.00045993328,0.00005797259,0.0008031947,0.0066975257,0.0007892416,0.0032624565,0.096758455,0.8796222,0.00008110014],"about_ca_topic_score_codex":0.013033777,"about_ca_topic_score_gemma":0.014936691,"teacher_disagreement_score":0.045718446,"about_ca_system_score_codex":0.0047241547,"about_ca_system_score_gemma":0.007560908,"threshold_uncertainty_score":0.15294349},"labels":[],"label_agreement":null},{"id":"W6945660656","doi":"10.25549/webster-c100-13769","title":"The Widening Divide, 1989-06","year":2021,"lang":"en","type":"dataset","venue":"University of Southern California Digital Library","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Commission; Newspaper; Poverty; Inequality; Economic Justice; Variety (cybernetics)","score_opus":0.006547402176054701,"score_gpt":0.17560799486344977,"score_spread":0.16906059268739507,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6945660656","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00095783686,0.00013201714,0.00002663425,0.0002367182,0.00004190088,0.000008496685,0.997874,0.000107305226,0.000615052],"genre_scores_gemma":[0.000977505,0.00008357404,0.00008394924,0.000038370013,0.000012270461,0.00003617829,0.9982388,0.00001573177,0.00051366835],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999014,0.0001377786,0.00013193347,0.00024837712,0.00030555218,0.00016223374],"domain_scores_gemma":[0.9979494,0.00028903453,0.00029601663,0.00032049092,0.0008541942,0.00029089366],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012242415,0.0010885767,0.00080017856,0.0032529659,0.0010308415,0.0028859,0.0018834925,0.0010032849,0.011234005],"category_scores_gemma":[0.005058143,0.00054434635,0.0006740543,0.008262762,0.0003423693,0.0011031743,0.001755205,0.0014498112,0.01960521],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000035608613,0.000020692145,0.005521691,0.00015042613,0.00001500322,0.000025005142,0.00004816536,0.00013067714,0.000022184639,0.00020138976,0.99141663,0.0024125413],"study_design_scores_gemma":[0.00024011909,0.000027878337,0.100370474,0.00038049373,0.0000424935,0.00009179094,0.0007000387,0.0009978967,0.00035134822,0.000687953,0.8960736,0.00003597843],"about_ca_topic_score_codex":0.12945923,"about_ca_topic_score_gemma":0.22351053,"teacher_disagreement_score":0.12945923,"about_ca_system_score_codex":0.0024203,"about_ca_system_score_gemma":0.0030829636,"threshold_uncertainty_score":0.25741142},"labels":[],"label_agreement":null},{"id":"W6945751204","doi":"10.25384/sage.25701862","title":"sj-tif-1-cjk-10.1177_20543581241238808 – Supplemental material for Pathways for Diagnosing and Treating CKD-Associated Pruritus: A Narrative Review","year":2024,"lang":"en","type":"other","venue":"Sage Journals Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Narrative review; Narrative; Kidney disease; MEDLINE; Public health; Alternative medicine","score_opus":0.05545420600520633,"score_gpt":0.3619169428006475,"score_spread":0.30646273679544117,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6945751204","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015512703,0.4200467,0.0076080062,0.060706448,0.012078624,0.001694706,0.21065888,0.0045868475,0.2810685],"genre_scores_gemma":[0.016278319,0.5380588,0.018739631,0.038285278,0.009243771,0.002856622,0.1665994,0.0020480938,0.20789008],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99916196,0.00022763545,0.00016131997,0.00007811443,0.00031685864,0.000054180167],"domain_scores_gemma":[0.9916425,0.0050813495,0.0008935707,0.00018999491,0.0016792075,0.000513367],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0019520657,0.0005409595,0.00089189253,0.0033748685,0.00039636088,0.002635087,0.0011796165,0.0017766502,0.34567395],"category_scores_gemma":[0.013483826,0.0002680345,0.0009983416,0.0038388679,0.00031341737,0.0016953795,0.001399539,0.001558039,0.10594774],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010608866,0.000021843103,0.00021026861,0.011936467,0.00006697172,0.00006624206,0.00004301808,0.00007144879,0.00024086963,0.0018798434,0.7820222,0.20333478],"study_design_scores_gemma":[0.000054233515,0.000021272184,0.00068396935,0.007040706,0.000090582005,0.00017417067,0.000028462695,0.00003541263,0.00016052915,0.0010742696,0.9906246,0.000011718063],"about_ca_topic_score_codex":0.003989168,"about_ca_topic_score_gemma":0.009124624,"teacher_disagreement_score":0.34567395,"about_ca_system_score_codex":0.0012499987,"about_ca_system_score_gemma":0.0033345972,"threshold_uncertainty_score":0.9333167},"labels":[],"label_agreement":null},{"id":"W6945782079","doi":"10.25384/sage.22285100","title":"sj-docx-1-nad-10.1177_14550725231160335 - Supplemental material for The public-private decision for alcohol retail systems: Examining the economic, health, and social impacts of alternative systems in Finland","year":2023,"lang":"en","type":"article","venue":"Sage Journals Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Victoria","funders":"","keywords":"Alcohol; Decision process; Retail sales; Alcohol consumption","score_opus":0.12690167470271185,"score_gpt":0.37563701795306975,"score_spread":0.2487353432503579,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6945782079","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00013225211,0.00003520558,0.00021518143,0.00032334396,0.00010976461,0.000047366062,0.98847497,0.0009778877,0.009683978],"genre_scores_gemma":[0.0034024252,0.00025489464,0.0024096915,0.0006549389,0.00015852462,0.0006091411,0.9466888,0.0035546492,0.042267077],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988122,0.00010976532,0.00017209488,0.00019138063,0.0005381614,0.00017638918],"domain_scores_gemma":[0.98260325,0.009471436,0.0009305908,0.00094417145,0.0049115843,0.0011388751],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016362811,0.0010291919,0.0010514496,0.0053181523,0.0016539728,0.005842469,0.0023375023,0.0022118052,0.90179634],"category_scores_gemma":[0.024326742,0.0010829573,0.00093716144,0.0103457235,0.000576501,0.00510007,0.0031668895,0.0016863662,0.6603543],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040460374,0.000028893644,0.00049555104,0.00060237653,0.0000071608447,0.000023613746,0.00009290585,0.000052390435,0.00007514272,0.0004979275,0.99282116,0.0052624936],"study_design_scores_gemma":[0.00028222761,0.000025642708,0.007767347,0.00096581236,0.000024169729,0.00009066013,0.00073327427,0.00018088667,0.0004870688,0.0027061934,0.98668057,0.00005621592],"about_ca_topic_score_codex":0.02395953,"about_ca_topic_score_gemma":0.03494287,"teacher_disagreement_score":0.90179634,"about_ca_system_score_codex":0.0023371764,"about_ca_system_score_gemma":0.004118637,"threshold_uncertainty_score":0.1400755},"labels":[],"label_agreement":null},{"id":"W6946022640","doi":"10.25384/sage.21754014","title":"Visual abstract – Supplemental material for A Systematic Review and Meta-Analysis of Robot-Assisted Mitral Valve Repair","year":2022,"lang":"en","type":"other","venue":"Sage Journals Data","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta","funders":"","keywords":"Mitral valve repair; Cardiothoracic surgery; Mitral valve; Mitral valve annuloplasty; Systematic review","score_opus":0.08441011945168175,"score_gpt":0.38751171196981,"score_spread":0.30310159251812824,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946022640","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004771727,0.0024208196,0.0012263319,0.0013274013,0.0004520015,0.0010472481,0.98646647,0.0010219238,0.005560729],"genre_scores_gemma":[0.025998732,0.012663202,0.029717026,0.008910929,0.0027230056,0.026885465,0.794754,0.0032196422,0.095128074],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983876,0.00037544093,0.00041920945,0.00024621948,0.00042869995,0.00014289377],"domain_scores_gemma":[0.9699836,0.023237247,0.0027738751,0.00068373204,0.0027971885,0.00052441057],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0030129005,0.0015313952,0.0024919654,0.005747761,0.0004151333,0.0023951696,0.0015723963,0.0014876802,0.8134258],"category_scores_gemma":[0.044827603,0.0007781197,0.003077963,0.0074267765,0.00024244978,0.0020858413,0.0017431554,0.0010380495,0.093911834],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00064598257,0.0000704685,0.001404675,0.18613727,0.0014540785,0.000100827565,0.00008641873,0.00040201572,0.00036308254,0.0017723157,0.77560776,0.031955145],"study_design_scores_gemma":[0.008236377,0.0004066586,0.020957956,0.069303036,0.0069611543,0.0007305238,0.00024241346,0.0017565148,0.00075404637,0.019701067,0.8707467,0.0002036397],"about_ca_topic_score_codex":0.0044306973,"about_ca_topic_score_gemma":0.014952359,"teacher_disagreement_score":0.8134258,"about_ca_system_score_codex":0.0014410355,"about_ca_system_score_gemma":0.00437211,"threshold_uncertainty_score":0.26612538},"labels":[],"label_agreement":null},{"id":"W6946155213","doi":"10.3389/fpsyt.2024.1286078.s001","title":"Data_Sheet_1_Prototyping the implementation of a suicide prevention protocol in primary care settings using PDSA cycles: a mixed method study.PDF","year":2024,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Suicide prevention; Protocol (science); Mental health; Poison control; PDCA; Primary care; Program evaluation; Addiction; Occupational safety and health","score_opus":0.053194342366173056,"score_gpt":0.42348002723880573,"score_spread":0.3702856848726327,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946155213","genre_codex":"protocol","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013493393,0.00034115915,0.010578385,0.0019008378,0.00043462962,0.7330754,0.21932414,0.0012743482,0.019577703],"genre_scores_gemma":[0.0062725106,0.00022854996,0.029745273,0.0006812506,0.00004035778,0.9479474,0.011291526,0.00013009526,0.0036630698],"study_design_codex":"design_other","study_design_gemma":"qualitative","domain_scores_codex":[0.96800107,0.019742673,0.0056681447,0.0012069953,0.004408783,0.0009723796],"domain_scores_gemma":[0.845679,0.10812089,0.006425179,0.010651955,0.026381813,0.002741199],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07598103,0.0010392863,0.0017811371,0.005088358,0.0024132954,0.003417867,0.003911976,0.0017510427,0.2592317],"category_scores_gemma":[0.1229672,0.0020474808,0.0025156275,0.006310778,0.0010759388,0.0024137832,0.002469888,0.0029632542,0.03632617],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0070426697,0.006650246,0.0065546124,0.035707727,0.00017057366,0.00027126636,0.010472134,0.0015429796,0.0007085705,0.005500798,0.45520198,0.47017634],"study_design_scores_gemma":[0.023666391,0.010121802,0.051007956,0.029830141,0.00028905002,0.00018555176,0.0170883,0.0025518406,0.0038072295,0.005620205,0.85541695,0.00041461916],"about_ca_topic_score_codex":0.0060627838,"about_ca_topic_score_gemma":0.00871339,"teacher_disagreement_score":0.2592317,"about_ca_system_score_codex":0.0053690863,"about_ca_system_score_gemma":0.019063966,"threshold_uncertainty_score":0.86721635},"labels":[],"label_agreement":null},{"id":"W6946163103","doi":"10.3389/fpsyt.2022.812965.s004","title":"Data_Sheet_4_HEARTSMAP-U: Adapting a Psychosocial Self-Screening and Resource Navigation Support Tool for Use by Post-secondary Students.xlsx","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Relevance (law); Adaptation (eye); Psychosocial; Resource (disambiguation); Process (computing); Mental health; Focus group","score_opus":0.032700665677286384,"score_gpt":0.3165680337427625,"score_spread":0.2838673680654761,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946163103","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00582111,0.0003204167,0.007478388,0.0051708193,0.0008445503,0.036753792,0.9073463,0.005588436,0.030676248],"genre_scores_gemma":[0.026660819,0.0015969835,0.08126038,0.0067828703,0.0005708901,0.39171585,0.41743547,0.0038339214,0.07014282],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99422014,0.0020939114,0.0016426094,0.00035341867,0.001442827,0.00024710863],"domain_scores_gemma":[0.9161929,0.055083044,0.0038632886,0.00430771,0.018835584,0.0017175216],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.015957579,0.0012420181,0.0016346377,0.0046787257,0.0013704149,0.0028710617,0.0022621213,0.0016811586,0.54500824],"category_scores_gemma":[0.07827648,0.001159036,0.0019534521,0.003543246,0.00070649,0.0033105782,0.002339162,0.0024240469,0.10579813],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008360498,0.00037718302,0.0034967747,0.0062325043,0.00004658413,0.00009748211,0.0009864791,0.0003671198,0.0002308393,0.0014514644,0.89323866,0.092638865],"study_design_scores_gemma":[0.0034974432,0.00055141776,0.034973323,0.009333701,0.000077075834,0.00018738478,0.0026484362,0.0007951244,0.001496557,0.0049937405,0.94122505,0.00022068564],"about_ca_topic_score_codex":0.004247954,"about_ca_topic_score_gemma":0.008190183,"teacher_disagreement_score":0.54500824,"about_ca_system_score_codex":0.002037534,"about_ca_system_score_gemma":0.0060606236,"threshold_uncertainty_score":0.6489905},"labels":[],"label_agreement":null},{"id":"W6946181520","doi":"10.26180/21655355","title":"The Legacy of a Villain: King John and the Rule of Law","year":2013,"lang":"en","type":"other","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Honour; Charter; Human rights; Relation (database); Rule of law","score_opus":0.020827041950392085,"score_gpt":0.2920126647373541,"score_spread":0.271185622786962,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946181520","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0071486146,0.018624632,0.0052074823,0.3432777,0.004042038,0.000027966194,0.00014304953,0.00012373402,0.62140477],"genre_scores_gemma":[0.16013943,0.009611172,0.0036897238,0.052601613,0.0015490999,0.000048302914,0.00010525975,0.00022972104,0.77202576],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9967962,0.0012508639,0.0001402336,0.00045726282,0.001007026,0.00034844343],"domain_scores_gemma":[0.99446976,0.0036962966,0.00025083427,0.00040561377,0.0006214232,0.00055598584],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004132145,0.00029524186,0.00034116997,0.00090297085,0.011722942,0.014291409,0.00084015564,0.004762402,0.008977416],"category_scores_gemma":[0.011592264,0.00027992393,0.00025128067,0.0012313073,0.016766023,0.008569404,0.003691802,0.007210494,0.0019573197],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000011474039,0.000010664303,0.00024525664,0.00004982817,0.0000033426084,0.00024070559,0.016094204,0.0000629952,0.000072679824,0.7017544,0.25817257,0.023281949],"study_design_scores_gemma":[0.0000016483024,0.0000024527355,0.00024308442,0.00012637154,0.0000016110841,0.00009484625,0.0051566963,0.000086962544,0.00009517742,0.06334462,0.9308365,0.000010024865],"about_ca_topic_score_codex":0.091503106,"about_ca_topic_score_gemma":0.18856922,"teacher_disagreement_score":0.091503106,"about_ca_system_score_codex":0.0072891153,"about_ca_system_score_gemma":0.011478062,"threshold_uncertainty_score":0.18194103},"labels":[],"label_agreement":null},{"id":"W6946220539","doi":"10.3389/fphys.2022.970016.s003","title":"Table2_Adiponectin, leptin, cortisol, neuropeptide Y and profile of mood states in athletes participating in an ultramarathon during winter: An observational study.docx","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Athletes; Mood; Affect (linguistics); Depressed mood; Observational study; Exertion; Neuropeptide Y receptor; Leptin","score_opus":0.09834966313339186,"score_gpt":0.3395230605269878,"score_spread":0.24117339739359592,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946220539","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24204628,0.0005620773,0.0008236661,0.00045483973,0.00020772975,0.0012503967,0.74986976,0.00018801894,0.00459724],"genre_scores_gemma":[0.6603637,0.0009776592,0.0040126177,0.0010154004,0.00032225825,0.006775107,0.29824847,0.00013714914,0.028147688],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99985874,0.000020209864,0.000036376612,0.00003532136,0.000022006854,0.000027364005],"domain_scores_gemma":[0.99948895,0.00015379072,0.00014814161,0.000032220152,0.00011536653,0.00006162453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002888787,0.0003825628,0.00056840957,0.0009804462,0.0006373097,0.00042999312,0.00059588056,0.0003594644,0.06622725],"category_scores_gemma":[0.0011873133,0.0002596034,0.0005026157,0.0014614244,0.00013918662,0.00051233283,0.00031344726,0.00057133473,0.0044131563],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022722378,0.0009912885,0.6873551,0.0038401054,0.00037554099,0.0009084641,0.00065134646,0.0004965019,0.0017342813,0.00031164606,0.2684472,0.03261637],"study_design_scores_gemma":[0.00024064019,0.0005114657,0.97368115,0.00040182754,0.00011026665,0.0007611712,0.00086433114,0.0005330543,0.00032912416,0.00018561816,0.0223413,0.000039908904],"about_ca_topic_score_codex":0.010285456,"about_ca_topic_score_gemma":0.012707863,"teacher_disagreement_score":0.06622725,"about_ca_system_score_codex":0.00041358275,"about_ca_system_score_gemma":0.00048104636,"threshold_uncertainty_score":0.22155225},"labels":[],"label_agreement":null},{"id":"W6946248185","doi":"10.34989/swp-2025-17","title":"Correcting Selection Bias in a Non-Probability Two-Phase Payment Survey","year":2025,"lang":"en","type":"article","venue":"Bank of Canada Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Bank of Canada","funders":"","keywords":"Selection (genetic algorithm); Selection bias; Calibration; Variance (accounting); Payment; Sample (material); Estimation; Sample size determination","score_opus":0.09431446954622,"score_gpt":0.4124850552878784,"score_spread":0.31817058574165835,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946248185","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02216116,0.000055088494,0.97553307,0.00041972453,0.000046532234,0.00047219597,0.0001968908,0.00022586691,0.0008893864],"genre_scores_gemma":[0.40899923,0.00012476304,0.5856552,0.00058950053,0.00010194669,0.0019539476,0.00065328786,0.0000829841,0.0018390868],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9361996,0.05144659,0.0020389522,0.0047315634,0.0047688605,0.00081433187],"domain_scores_gemma":[0.7043849,0.24360417,0.01513969,0.027264146,0.008849874,0.0007572604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.086092524,0.00067650276,0.0012874268,0.0017514036,0.0010551206,0.002437114,0.0032002132,0.0019256154,0.0051180273],"category_scores_gemma":[0.34626788,0.000987687,0.0014110229,0.0031146826,0.002062035,0.004215858,0.0032005538,0.002127115,0.0007919416],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007692753,0.0005312889,0.12908953,0.00081638584,0.0008590406,0.00062542653,0.0027940692,0.09547309,0.0026656936,0.4004605,0.008107005,0.35780865],"study_design_scores_gemma":[0.00038919368,0.00065728667,0.034763664,0.00026229458,0.00028293955,0.0005381342,0.00052373286,0.59375787,0.004332909,0.3513634,0.013009658,0.00011882909],"about_ca_topic_score_codex":0.0037709163,"about_ca_topic_score_gemma":0.0036419965,"teacher_disagreement_score":0.086092524,"about_ca_system_score_codex":0.0011436592,"about_ca_system_score_gemma":0.0028923852,"threshold_uncertainty_score":0.45530623},"labels":[],"label_agreement":null},{"id":"W6946250206","doi":"10.26180/19246038","title":"The contemporary role of shareholder ratification and authorisation of breaches of director’s duties","year":2022,"lang":"en","type":"dissertation","venue":"Monash University","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ratification; Statutory law; Doctrine; Authorization; Shareholder; Limiting","score_opus":0.01238782197379408,"score_gpt":0.2258951832064085,"score_spread":0.21350736123261443,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946250206","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0821478,0.017176677,0.10215776,0.098833166,0.0020457502,0.000097654534,0.00017106048,0.00033984147,0.6970303],"genre_scores_gemma":[0.9444969,0.005982585,0.011976256,0.005504141,0.0019434607,0.00011405307,0.00011115357,0.00019141233,0.029680045],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9602203,0.017715292,0.0024839453,0.006420243,0.010705064,0.002455237],"domain_scores_gemma":[0.89767534,0.069282606,0.009287691,0.014012078,0.007280218,0.0024620732],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0451315,0.00035136438,0.0006811874,0.004043247,0.007906296,0.02187381,0.0029238944,0.00668458,0.007344446],"category_scores_gemma":[0.082826935,0.0006897085,0.00058120606,0.004676608,0.0587459,0.03277497,0.0075658225,0.008485011,0.0017660712],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000128862275,0.000009432083,0.00066419883,0.000038230584,0.0000039283614,0.000037905666,0.0043405504,0.00015207056,0.00014944073,0.975426,0.0017155884,0.017449854],"study_design_scores_gemma":[0.00003014731,0.000033575612,0.0038642122,0.00064305274,0.000031183383,0.00031092562,0.0061315573,0.0020047307,0.0021804285,0.6528616,0.33183283,0.000075746466],"about_ca_topic_score_codex":0.008071763,"about_ca_topic_score_gemma":0.0046768324,"teacher_disagreement_score":0.0451315,"about_ca_system_score_codex":0.011985591,"about_ca_system_score_gemma":0.011611692,"threshold_uncertainty_score":0.23868108},"labels":[],"label_agreement":null},{"id":"W6946251731","doi":"10.34630/sensos-e.v12i2.5921","title":"Canadian Museum for Human Rights: Uma oportunidade para o aprendizado de ética","year":2024,"lang":"pt","type":"article","venue":"Instituto Politécnico do Porto","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Manitoba; University of Winnipeg","funders":"","keywords":"Context (archaeology); Quality (philosophy); Human being; Action (physics); Human life","score_opus":0.03363087452146024,"score_gpt":0.3326972705588362,"score_spread":0.29906639603737595,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946251731","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.02882527,0.026500545,0.019169454,0.27805674,0.004397662,0.00031640948,0.0024059978,0.0006310451,0.6396969],"genre_scores_gemma":[0.56081855,0.03181062,0.03823917,0.027077453,0.0011019731,0.0002652194,0.0018654385,0.0007537146,0.33806792],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"qualitative","domain_scores_codex":[0.9912042,0.0015132208,0.00026293023,0.0009565833,0.0046529514,0.0014101319],"domain_scores_gemma":[0.9851483,0.0030667612,0.0006712046,0.0017846529,0.006159387,0.0031696241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009349329,0.0006232481,0.000713441,0.007201171,0.027284924,0.020797605,0.003278748,0.0036993306,0.027244216],"category_scores_gemma":[0.017197473,0.0005447647,0.00067956693,0.01099882,0.026629936,0.008437987,0.00967886,0.005187483,0.0017939566],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003488118,0.000027286487,0.00557902,0.00048090218,0.000030003574,0.00048450363,0.029795395,0.00026532283,0.00076745125,0.65587866,0.17577435,0.13088222],"study_design_scores_gemma":[0.000005927747,0.000006430258,0.006431715,0.00052496634,0.000019288262,0.00019450382,0.015631609,0.00019964571,0.00027003177,0.025248876,0.9514158,0.00005137174],"about_ca_topic_score_codex":0.94445926,"about_ca_topic_score_gemma":0.9713566,"teacher_disagreement_score":0.07389517,"about_ca_system_score_codex":0.07389517,"about_ca_system_score_gemma":0.23933819,"threshold_uncertainty_score":0.5361495},"labels":[],"label_agreement":null},{"id":"W6946266929","doi":"10.25921/6w5n-b680","title":"NOAA/WDS Paleoclimatology - Luckman - Highwood Pass - PCEN - ITRDB CANA438","year":2014,"lang":"en","type":"dataset","venue":"National Oceanic and Atmospheric Administration (NOAA) National Centers for Environmental Information (NCEI)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Chronology; Principal component analysis; Paleoclimatology; Population; Calibration; Period (music); Residual; Range (aeronautics); Confidence interval","score_opus":0.008386728690912204,"score_gpt":0.24916864006751763,"score_spread":0.24078191137660543,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946266929","genre_codex":"other","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019726,0.0034087093,0.019233963,0.0030371258,0.0013655183,0.000670175,0.17884625,0.006308243,0.7674039],"genre_scores_gemma":[0.047985166,0.002539334,0.04752357,0.0005604152,0.00013802583,0.00036973014,0.15408218,0.0046452475,0.7421564],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993693,0.00005256531,0.000022809594,0.00012858283,0.00034969105,0.00007713627],"domain_scores_gemma":[0.9988211,0.000047565743,0.000052598534,0.00015210576,0.00074445584,0.0001821313],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012340136,0.00076645083,0.0006816559,0.0022401486,0.0018186797,0.0029457803,0.0013635688,0.00054629263,0.15618004],"category_scores_gemma":[0.0010515081,0.0006958051,0.00043646756,0.0024439155,0.0004891677,0.00092465075,0.0013441113,0.0007960917,0.04764062],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00047790544,0.00018178062,0.022048809,0.00044882795,0.00008993817,0.0002483597,0.0003345007,0.00523943,0.010743529,0.0147745395,0.55703735,0.3883751],"study_design_scores_gemma":[0.000052341078,0.000018038343,0.026449995,0.000119850156,0.000017041622,0.00003902991,0.00010639543,0.0037054361,0.0020766836,0.001252437,0.96612877,0.00003400972],"about_ca_topic_score_codex":0.8355313,"about_ca_topic_score_gemma":0.9143803,"teacher_disagreement_score":0.1644687,"about_ca_system_score_codex":0.006969671,"about_ca_system_score_gemma":0.016740747,"threshold_uncertainty_score":0.5224743},"labels":[],"label_agreement":null},{"id":"W6946361347","doi":"10.3389/finsc.2023.1104793.s003","title":"Table_1_Revisiting fall armyworm population movement in the United States and Canada.docx","year":2023,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Biological dispersal; Gene flow; Overwintering; Range (aeronautics); Population; Latitude","score_opus":0.026266006096584818,"score_gpt":0.27377942373062675,"score_spread":0.24751341763404194,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946361347","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00091352424,0.0000796346,0.00009288332,0.00012178499,0.000039379,0.000031952888,0.99523336,0.00012665303,0.0033608025],"genre_scores_gemma":[0.012187385,0.0005699599,0.0017767538,0.000336457,0.00003111589,0.00017291939,0.9704158,0.0002667858,0.014242817],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996331,0.000016730814,0.000027029948,0.000060141985,0.0001686098,0.00009436403],"domain_scores_gemma":[0.9951988,0.0003097128,0.0001513799,0.00014376442,0.0039822124,0.00021415383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00045138676,0.00092439126,0.0007016808,0.0050456487,0.0020867505,0.0017163082,0.002109191,0.00040428757,0.17129907],"category_scores_gemma":[0.004279281,0.00048241447,0.0011712511,0.01263861,0.0003005361,0.0009605917,0.0010252377,0.0008927504,0.022397902],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026494345,0.000014582546,0.020279568,0.00047713183,0.000049439735,0.000043853175,0.00009699263,0.00085530785,0.0000574229,0.0005914042,0.9683747,0.009133155],"study_design_scores_gemma":[0.00018354999,0.000019426338,0.2513807,0.0011459203,0.00011948257,0.0001556763,0.001565286,0.002904238,0.0005011629,0.001064741,0.74087185,0.000087938664],"about_ca_topic_score_codex":0.98391116,"about_ca_topic_score_gemma":0.98793495,"teacher_disagreement_score":0.17129907,"about_ca_system_score_codex":0.009593681,"about_ca_system_score_gemma":0.026044495,"threshold_uncertainty_score":0.57305247},"labels":[],"label_agreement":null},{"id":"W6946433301","doi":"10.3389/fmars.2024.1462905.s001","title":"Image1_Transcriptomic responses to hypoxia in two populations of eastern oyster with differing tolerance.jpeg","year":2024,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Hypoxia (environmental); Crassostrea; Oyster; Gene; Threatened species; Gene expression; Eastern oyster; Habitat","score_opus":0.04395961367108449,"score_gpt":0.31993036462534014,"score_spread":0.27597075095425566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946433301","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.3498115,0.0011387698,0.008185068,0.0007149179,0.0002822925,0.00014304837,0.61649626,0.0032399907,0.01998819],"genre_scores_gemma":[0.36245576,0.0014857294,0.029656416,0.0007583013,0.000103136044,0.00059647067,0.56416327,0.0013897857,0.039391078],"study_design_codex":"bench_or_experimental","study_design_gemma":"not_applicable","domain_scores_codex":[0.99991596,0.0000027414535,0.0000049275295,0.000034156765,0.00002430625,0.000017962482],"domain_scores_gemma":[0.99985254,0.000045962577,0.000025917172,0.000012122253,0.000042081367,0.000021465077],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.000110843794,0.00027021702,0.00026705957,0.0009278127,0.00035725662,0.00048146222,0.0002559769,0.00041009882,0.021301309],"category_scores_gemma":[0.0002409397,0.00015933311,0.00035083285,0.0014030568,0.00017522753,0.00029054063,0.00037245816,0.0003794542,0.0028421797],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012879618,0.00015512726,0.049555395,0.0021348184,0.00016455485,0.0009162203,0.000853425,0.0012421074,0.77125174,0.0013242627,0.078525804,0.092588626],"study_design_scores_gemma":[0.000119850054,0.00021633087,0.7882412,0.00017758104,0.00017429468,0.0012679499,0.0008549134,0.0073451516,0.0870131,0.0009135999,0.11358282,0.00009316998],"about_ca_topic_score_codex":0.008226249,"about_ca_topic_score_gemma":0.013807591,"teacher_disagreement_score":0.9786987,"about_ca_system_score_codex":0.00027710208,"about_ca_system_score_gemma":0.0002912001,"threshold_uncertainty_score":0.071259916},"labels":[],"label_agreement":null},{"id":"W6946462341","doi":"10.3389/fvets.2019.00023.s001","title":"Table_1_Military Veterans and Their PTSD Service Dogs: Associations Between Training Methods, PTSD Severity, Dog Behavior, and the Human-Animal Bond.xlsx","year":2019,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Closeness; Population; Service member; Intervention (counseling); Military service; Service (business)","score_opus":0.09777010872371315,"score_gpt":0.3723701410007837,"score_spread":0.2746000322770705,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946462341","genre_codex":"empirical","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.55871063,0.0015186074,0.00048735642,0.004496403,0.0003793764,0.00039066692,0.3900757,0.0002834956,0.043657824],"genre_scores_gemma":[0.840303,0.0017947543,0.0014868652,0.0014595188,0.0003203524,0.00078727177,0.108948134,0.00007611678,0.04482394],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99988616,0.000020199359,0.000020134466,0.000021525904,0.00002750641,0.00002437907],"domain_scores_gemma":[0.99884254,0.00041613093,0.0003995895,0.00002462275,0.00017204054,0.00014512683],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00015982473,0.00025648554,0.0003137253,0.00074465707,0.0006290936,0.00034102864,0.0003661507,0.00047677214,0.16103402],"category_scores_gemma":[0.0016771882,0.00012815133,0.00026351254,0.0009621315,0.00011789537,0.0008058544,0.00028358473,0.00044348175,0.0067346],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00043960699,0.00042942903,0.7623746,0.00053630205,0.000074340765,0.00027665813,0.00046528564,0.00022393672,0.00041747768,0.00042137824,0.21116927,0.023171723],"study_design_scores_gemma":[0.000061375256,0.00031753985,0.97908807,0.0002137562,0.000017491777,0.00047414616,0.0018511456,0.0003817986,0.00011233544,0.00025765557,0.017207759,0.00001705269],"about_ca_topic_score_codex":0.0075755203,"about_ca_topic_score_gemma":0.010381513,"teacher_disagreement_score":0.16103402,"about_ca_system_score_codex":0.00022625935,"about_ca_system_score_gemma":0.0003473849,"threshold_uncertainty_score":0.5387125},"labels":[],"label_agreement":null},{"id":"W6946583225","doi":"10.3389/fpsyt.2022.1015443.s001","title":"Table_1_Promoting self-change in cannabis use disorder: Findings from a randomized trial.DOCX","year":2022,"lang":"en","type":"dataset","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cannabis; Abstinence; Gee; Randomized controlled trial; Motivational interviewing; Psychological intervention; Workbook; Cannabis Dependence","score_opus":0.037695064381328025,"score_gpt":0.28955547670963155,"score_spread":0.25186041232830353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946583225","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.03653817,0.024714993,0.005637806,0.013572222,0.00611535,0.30436304,0.53201663,0.0031269821,0.07391489],"genre_scores_gemma":[0.16177514,0.02163976,0.02496609,0.016234875,0.002549728,0.6013986,0.07010758,0.0009020676,0.100426175],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983204,0.00093677716,0.00021808936,0.00015071583,0.00023581262,0.0001382089],"domain_scores_gemma":[0.99160236,0.0062287906,0.00088550954,0.0002163381,0.0006643365,0.00040269998],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0056611965,0.0011500973,0.0019824423,0.0010298255,0.00057788985,0.0013624331,0.0015886939,0.0015634111,0.43710673],"category_scores_gemma":[0.017030824,0.00062646874,0.002912972,0.0012240858,0.00034231448,0.0017300987,0.00063533353,0.0019647616,0.016203146],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.28724995,0.003967658,0.0016296082,0.12038192,0.0039430396,0.00013598887,0.00017415857,0.0008401363,0.0009436245,0.0033214788,0.40024936,0.1771631],"study_design_scores_gemma":[0.84814864,0.015256049,0.013305838,0.019804789,0.004698127,0.00011925479,0.00022925748,0.0011319296,0.0011363967,0.0029603727,0.093114436,0.000094766445],"about_ca_topic_score_codex":0.0021228245,"about_ca_topic_score_gemma":0.0051111463,"teacher_disagreement_score":0.43710673,"about_ca_system_score_codex":0.0011475475,"about_ca_system_score_gemma":0.0019085949,"threshold_uncertainty_score":0.80289894},"labels":[],"label_agreement":null},{"id":"W6946821065","doi":"10.34943/b53ee73a-b2d1-4368-8103-f4310b7f4211","title":"Hecate Strait Hydrophone Deployed 2014-06-30","year":2017,"lang":"en","type":"dataset","venue":"Ocean Networks Canada Society","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Hydrophone; Underwater; Data logger; Waveform; Sound (geography); Software deployment","score_opus":0.008027624169892873,"score_gpt":0.23906022585956205,"score_spread":0.2310326016896692,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6946821065","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015809551,0.000094593925,0.00026721944,0.00012421729,0.00004311432,0.00002623834,0.9956033,0.00061429327,0.0016460605],"genre_scores_gemma":[0.0014278,0.00004352064,0.0003518419,0.000028003822,0.00000511142,0.000035409772,0.99723166,0.0000384836,0.00083814136],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99909484,0.00010237379,0.00008367148,0.00021411279,0.00035249858,0.00015247076],"domain_scores_gemma":[0.99867886,0.00017600991,0.000103814265,0.00028954772,0.00058873795,0.00016301496],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00097634626,0.0016666029,0.0008607595,0.0026552887,0.00075733016,0.0012564133,0.0020446894,0.0011917175,0.01052055],"category_scores_gemma":[0.003171231,0.00034068275,0.0006601369,0.0039002893,0.0005690378,0.0010627249,0.0014808775,0.0012790273,0.023335334],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001313435,0.0000544419,0.004123781,0.00041745478,0.000041924177,0.000101144804,0.000072922565,0.0011109861,0.0004589089,0.00092598953,0.98642886,0.006132157],"study_design_scores_gemma":[0.0001178705,0.000033119966,0.01440713,0.00023020816,0.000030176989,0.0001518795,0.00032070914,0.0022943711,0.0013603853,0.0013744691,0.97963125,0.000048359194],"about_ca_topic_score_codex":0.09550804,"about_ca_topic_score_gemma":0.20192571,"teacher_disagreement_score":0.90449196,"about_ca_system_score_codex":0.0021717956,"about_ca_system_score_gemma":0.0037001919,"threshold_uncertainty_score":0.18990421},"labels":[],"label_agreement":null},{"id":"W6947083180","doi":"10.3897/zookeys.573.7878.figures86-89","title":"Figures 86-89 from: Klimaszewski J, Webster RP, Langor DW, Sikes D, Bourdon C, Godin B, Ernst C (2016) A review of Canadian and Alaskan species of the genus Liogluta Thomson, and descriptions of three new species (Coleoptera, Staphylinidae, Aleocharinae). In: Webster RP, Bouchard P, Klimaszewski J (Eds) The Coleoptera of New Brunswick and Canada: providing baseline biodiversity and natural history data. ZooKeys 573: 217–256. https://doi.org/10.3897/zookeys.573.7878","year":2016,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Aedeagus; Habitus; Dorsum; Genus","score_opus":0.04357707441342847,"score_gpt":0.22415274615308162,"score_spread":0.18057567173965314,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947083180","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00054483674,0.0017131628,0.0017705829,0.0007658837,0.0028690924,0.00019575971,0.08749621,0.0020048195,0.9026396],"genre_scores_gemma":[0.006223802,0.0030765515,0.0042786724,0.0005201595,0.000585771,0.00016366622,0.089035556,0.0020330013,0.8940827],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99969196,0.00001968938,0.00002284618,0.000051746134,0.00017769958,0.00003609852],"domain_scores_gemma":[0.999143,0.00016769556,0.000073333315,0.00008163426,0.0004131324,0.0001210943],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00023799784,0.0012719368,0.00065627554,0.002534776,0.0011355056,0.001745836,0.0010853332,0.0007231661,0.8517474],"category_scores_gemma":[0.0023361482,0.00032318558,0.0005134162,0.004013559,0.0006275265,0.0024024427,0.0011948853,0.0011709771,0.67083037],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000013765135,0.0000059859176,0.000109714936,0.00015761012,0.0000019496638,0.00004172769,0.000045888206,0.0000427934,0.00013053186,0.00086557365,0.9762383,0.022346059],"study_design_scores_gemma":[0.0000015312642,0.0000013819262,0.0004111806,0.00004970254,0.0000010825346,0.000069413436,0.000022091006,0.000013868925,0.000043715296,0.00022437511,0.99916005,0.0000016566074],"about_ca_topic_score_codex":0.022483338,"about_ca_topic_score_gemma":0.046774585,"teacher_disagreement_score":0.97751665,"about_ca_system_score_codex":0.0013998676,"about_ca_system_score_gemma":0.0015297113,"threshold_uncertainty_score":0.21146429},"labels":[],"label_agreement":null},{"id":"W6947622155","doi":"10.4224/crm.2025.hisn-1","title":"HISN-1: High Purity Tin Certified Reference Material for Tin Mass Fraction and Elemental Impurities","year":2025,"lang":"en","type":"other","venue":"NRC Digital Repository","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Tin; Mass fraction; Impurity; Certified reference materials; Fraction (chemistry)","score_opus":0.009909012412072012,"score_gpt":0.24609482522326276,"score_spread":0.23618581281119075,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947622155","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014049657,0.0028919908,0.15002473,0.00089331955,0.00057925825,0.00075366296,0.44363177,0.055873495,0.33130214],"genre_scores_gemma":[0.03360465,0.0029250544,0.16323975,0.0008007762,0.00014254761,0.00066430634,0.5821426,0.017481664,0.19899864],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9975484,0.00024846583,0.00015976069,0.0003915214,0.0015349272,0.00011696247],"domain_scores_gemma":[0.9973687,0.00034170493,0.00024139138,0.00054510223,0.0014013075,0.00010182646],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030182162,0.001398294,0.0008874738,0.0064556827,0.00092072453,0.0021554325,0.0027344015,0.001306478,0.06149532],"category_scores_gemma":[0.0051906286,0.0004330608,0.0005982929,0.005833384,0.0005836883,0.002028389,0.0021466408,0.00088174443,0.08234207],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000708128,0.00024793492,0.003965767,0.0028473618,0.00009958752,0.0006101959,0.00039927068,0.00160057,0.067917414,0.034800146,0.6275507,0.2592529],"study_design_scores_gemma":[0.000027850136,0.000022314067,0.0021141914,0.00018435954,0.000029373976,0.00030571278,0.000063520805,0.0014363646,0.04401298,0.004003748,0.9477563,0.00004330231],"about_ca_topic_score_codex":0.0063143377,"about_ca_topic_score_gemma":0.008591527,"teacher_disagreement_score":0.06149532,"about_ca_system_score_codex":0.0013244484,"about_ca_system_score_gemma":0.0018840814,"threshold_uncertainty_score":0.20572233},"labels":[],"label_agreement":null},{"id":"W6947655074","doi":"10.48448/7z21-dr51","title":"Don’t Trust ChatGPT when your Question is not in English: A Study of Multilingual Abilities and Types of LLMs","year":2023,"lang":"en","type":"other","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia; University of Alberta","funders":"","keywords":"Variety (cybernetics); Generalization; Natural language; Natural (archaeology); Phenomenon; Variation (astronomy)","score_opus":0.04203779765758684,"score_gpt":0.34329207202300854,"score_spread":0.3012542743654217,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947655074","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9764846,0.00024438326,0.010937463,0.0005006296,0.00002815248,0.000095683434,0.00039423353,0.00043291232,0.010881935],"genre_scores_gemma":[0.9936412,0.00010749785,0.004329781,0.00018844685,0.00001605983,0.000041255462,0.00042628607,0.00012655898,0.0011229499],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9921508,0.0045956713,0.00057476247,0.0013561174,0.000989368,0.00033329957],"domain_scores_gemma":[0.8672468,0.1096098,0.006991022,0.010094248,0.00418891,0.0018692134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01012275,0.00057904515,0.0005393511,0.0016526561,0.000770263,0.002365892,0.0012310416,0.00093316747,0.0034189927],"category_scores_gemma":[0.09629865,0.00034701978,0.00046062696,0.0013382559,0.0019568894,0.007780595,0.0035921908,0.0016613093,0.00094898057],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0030664857,0.001555393,0.483888,0.0017237134,0.0007047645,0.0030611267,0.07448362,0.012176742,0.018858526,0.016777337,0.0072327177,0.37647158],"study_design_scores_gemma":[0.00028960028,0.0024585198,0.4407867,0.0006685368,0.00085813337,0.0070915516,0.06485087,0.31680754,0.03873601,0.08653115,0.04030529,0.0006161552],"about_ca_topic_score_codex":0.0034969952,"about_ca_topic_score_gemma":0.003458554,"teacher_disagreement_score":0.01012275,"about_ca_system_score_codex":0.0007360412,"about_ca_system_score_gemma":0.00083120255,"threshold_uncertainty_score":0.053534865},"labels":[],"label_agreement":null},{"id":"W6947658243","doi":"10.3897/zookeys.147.2098.figure3","title":"Figure 3 from: Bergeron C, Spence J, Volney J (2011) Landscape patterns of species-level association between ground-beetles and overstory trees in boreal forests of western Canada (Coleoptera, Carabidae). ZooKeys 147: 577-600. https://doi.org/10.3897/zookeys.147.2098","year":2011,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ordination; Taiga; Boreal; Basal area; Association (psychology); Dendrochronology","score_opus":0.03419737104176688,"score_gpt":0.22843041464909317,"score_spread":0.1942330436073263,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947658243","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0014887595,0.00034481735,0.0037007104,0.00083428813,0.00067438924,0.0001461082,0.5151804,0.013193138,0.46443745],"genre_scores_gemma":[0.009683488,0.0008088616,0.007940691,0.00042512818,0.00015613766,0.00016323684,0.37979284,0.008827037,0.5922024],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99977154,0.000011378478,0.000011079706,0.00004227279,0.00012370585,0.000040010782],"domain_scores_gemma":[0.99905294,0.00016563936,0.000049461272,0.00011459047,0.0004468217,0.00017054498],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00032803277,0.00096133974,0.00063244754,0.0021629713,0.000921816,0.0026111295,0.0009271986,0.0007199576,0.8243579],"category_scores_gemma":[0.0021826925,0.00038102412,0.0005615671,0.004096461,0.00040412036,0.0017505942,0.0014191577,0.00094840984,0.62252724],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020784833,0.000010596645,0.00034054532,0.00012979661,0.0000032567875,0.000038872862,0.00006254815,0.000048686263,0.00012127645,0.00042416732,0.98522645,0.0135730095],"study_design_scores_gemma":[0.000016270029,0.0000043693867,0.003512559,0.00008965917,0.0000051708485,0.000077769495,0.0000749492,0.000088420114,0.00024306386,0.0005229883,0.9953571,0.000007687661],"about_ca_topic_score_codex":0.04527409,"about_ca_topic_score_gemma":0.13416483,"teacher_disagreement_score":0.9547259,"about_ca_system_score_codex":0.001132944,"about_ca_system_score_gemma":0.0018432628,"threshold_uncertainty_score":0.2505321},"labels":[],"label_agreement":null},{"id":"W6947706383","doi":"10.3886/e194845","title":"Data and Code for: Careers and Intergenerational Income Mobility","year":2024,"lang":"en","type":"dataset","venue":"ICPSR Data Holdings","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary; Université du Québec à Montréal","funders":"","keywords":"Microdata (statistics); Census; Persistence (discontinuity); Survey of Income and Program Participation; Socioeconomic status; Social mobility; Occupational mobility","score_opus":0.06372702490498836,"score_gpt":0.3516261895954946,"score_spread":0.28789916469050625,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947706383","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006831191,0.000037216723,0.000106887055,0.00008784739,0.00001573802,0.000009573918,0.99828017,0.00012712427,0.00065228395],"genre_scores_gemma":[0.0012526758,0.00003617894,0.0004652515,0.000037143152,0.0000046833024,0.000054698685,0.9974268,0.000025024476,0.000697526],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99927956,0.00011402066,0.00010514047,0.00018537426,0.00020734478,0.00010848709],"domain_scores_gemma":[0.99785703,0.0005494542,0.00040204043,0.00046419725,0.00049527996,0.00023197326],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000768009,0.0010565727,0.0005725414,0.002869278,0.0006466977,0.0013696301,0.0013290219,0.0013017564,0.029411042],"category_scores_gemma":[0.0046784636,0.00038596976,0.0007002485,0.0063570915,0.0003239775,0.00080075517,0.0011734924,0.0014296257,0.03200082],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007366149,0.000052245625,0.008972955,0.00040593682,0.000029223493,0.00007390255,0.00008148842,0.0007029871,0.00020922738,0.0012479748,0.9824714,0.005678889],"study_design_scores_gemma":[0.00013463537,0.000017432441,0.029785039,0.00021985715,0.000018163018,0.00014486576,0.00025409873,0.0012745992,0.0005195343,0.0017234484,0.9658751,0.000033104756],"about_ca_topic_score_codex":0.03417993,"about_ca_topic_score_gemma":0.067635715,"teacher_disagreement_score":0.03417993,"about_ca_system_score_codex":0.0013186334,"about_ca_system_score_gemma":0.0017667069,"threshold_uncertainty_score":0.098389745},"labels":[],"label_agreement":null},{"id":"W6947757790","doi":"10.48448/3p62-rp66","title":"Responsible AI Considerations in Text Summarization Research: A Review of Current Practices | VIDEO","year":2023,"lang":"en","type":"other","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Automatic summarization; Task (project management); Focus (optics); Work (physics); Systematic review; Best practice","score_opus":0.31988270940585994,"score_gpt":0.5090560475914591,"score_spread":0.18917333818559912,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947757790","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0034747203,0.8692011,0.014231939,0.078413084,0.004226134,0.00046566033,0.0008501963,0.0004186561,0.028718563],"genre_scores_gemma":[0.030986905,0.9200616,0.02651884,0.0116220405,0.003991904,0.00088332425,0.0011939323,0.00039496378,0.004346477],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9773207,0.011467506,0.0045073037,0.0013662548,0.004908861,0.00042942495],"domain_scores_gemma":[0.6654394,0.2942148,0.010523724,0.0041643837,0.023614427,0.0020432442],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.043616433,0.0006594361,0.0009189173,0.022290852,0.0018881345,0.009749082,0.0023961586,0.003401597,0.011469213],"category_scores_gemma":[0.155785,0.0008207142,0.00096523843,0.019009942,0.0053678514,0.014078823,0.0035875856,0.003085157,0.0030100686],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008834248,0.000057862537,0.00089840667,0.03835737,0.000075888885,0.00019290143,0.008067783,0.00023633965,0.0006129833,0.019103946,0.065917395,0.8663909],"study_design_scores_gemma":[0.000026870079,0.0000649999,0.002423717,0.118344076,0.00014334296,0.0004734673,0.009043404,0.00039028586,0.00072138885,0.015315355,0.8529756,0.0000775305],"about_ca_topic_score_codex":0.007842224,"about_ca_topic_score_gemma":0.012306622,"teacher_disagreement_score":0.9563836,"about_ca_system_score_codex":0.0064667636,"about_ca_system_score_gemma":0.009729357,"threshold_uncertainty_score":0.23066849},"labels":[],"label_agreement":null},{"id":"W6947999492","doi":"10.48448/chfw-wa16","title":"ChatRetriever: Adapting Large Language Models for Generalized and Robust Conversational Dense Retrieval","year":2024,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Montréal","funders":"","keywords":"Robustness (evolution); Language model; Generalization; Rewriting; Session (web analytics); Interpretation (philosophy); Training set","score_opus":0.03032577417322737,"score_gpt":0.3014703699196813,"score_spread":0.27114459574645394,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6947999492","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.022640314,0.0006978761,0.94924414,0.00022119068,0.00008720602,0.00019221741,0.00047362747,0.024819594,0.0016238055],"genre_scores_gemma":[0.36839685,0.00041390018,0.61709815,0.00066628796,0.00014688687,0.0005079139,0.0029295557,0.0023349605,0.0075055254],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99856037,0.000527099,0.000085142754,0.00043874365,0.00027562532,0.000113063456],"domain_scores_gemma":[0.9967989,0.001568402,0.00013131993,0.0010119539,0.00036597418,0.00012350151],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002115206,0.0013992072,0.0013440517,0.00087774254,0.000553416,0.0014089161,0.0029773447,0.0014197904,0.003671675],"category_scores_gemma":[0.0082729,0.0006101211,0.0011989418,0.00064235914,0.00076652976,0.003522461,0.0030309297,0.0023179937,0.0034325258],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00056424615,0.00047691114,0.0017934896,0.00053228147,0.00030029757,0.0002840515,0.0007124599,0.1540642,0.053449634,0.0076752882,0.016691383,0.76345575],"study_design_scores_gemma":[0.000035380675,0.00011749709,0.00022851772,0.000012043295,0.00003366014,0.00009928558,0.0001042585,0.97574675,0.01058221,0.00830075,0.0046900087,0.000049690938],"about_ca_topic_score_codex":0.0074616536,"about_ca_topic_score_gemma":0.012299072,"teacher_disagreement_score":0.0074616536,"about_ca_system_score_codex":0.00075511413,"about_ca_system_score_gemma":0.0014368889,"threshold_uncertainty_score":0.01483649},"labels":[],"label_agreement":null},{"id":"W6948008966","doi":"10.48448/x6j7-0j45","title":"Ultrasonic vocalization analysis as a novel metric to assess home cage welfare in rats","year":2021,"lang":"en","type":"other","venue":"Underline Science Inc.","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Prince Edward Island","funders":"","keywords":"Cage; Agonistic behaviour; Animal welfare; Replicate; Metric (unit); Captivity; Welfare; Environmental enrichment","score_opus":0.032051037210935274,"score_gpt":0.33412956264465893,"score_spread":0.30207852543372365,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948008966","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95980376,0.0017620027,0.03395512,0.00013596036,0.0000948513,0.00023866791,0.0011338316,0.00045079872,0.0024249705],"genre_scores_gemma":[0.9336808,0.002168438,0.056920666,0.00030803276,0.000118488875,0.0011023964,0.0013152554,0.0001623739,0.0042234757],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99933535,0.00012755838,0.000051597905,0.0001578554,0.00023951368,0.00008807158],"domain_scores_gemma":[0.9990632,0.00013498797,0.00035782187,0.00009388533,0.0001932861,0.00015690301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00068789674,0.00078379834,0.0003963734,0.0015637038,0.00026028816,0.00047324583,0.00042184847,0.0005754571,0.0013432468],"category_scores_gemma":[0.00066675886,0.00023535157,0.00039777107,0.00056895154,0.00052834867,0.00046913925,0.00042955574,0.00076635997,0.00029655814],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005683648,0.00028861058,0.02206072,0.00016188752,0.00007593574,0.00008509219,0.00021534618,0.00018213381,0.95641,0.00011461714,0.00018791284,0.019649364],"study_design_scores_gemma":[0.000049095266,0.0104020275,0.6359718,0.00007128525,0.00037540784,0.0006900567,0.00085655274,0.002403749,0.3461369,0.00037247926,0.0025611932,0.00010955706],"about_ca_topic_score_codex":0.00078699,"about_ca_topic_score_gemma":0.0021480394,"teacher_disagreement_score":0.0015637038,"about_ca_system_score_codex":0.0002552409,"about_ca_system_score_gemma":0.00031374028,"threshold_uncertainty_score":0.004493594},"labels":[],"label_agreement":null},{"id":"W6948047220","doi":"10.48336/9aea-q877","title":"Green chemistry and an ocean based biorefinery approach for the valorization of Newfoundland and Labrador snow crab (Chionoecetes opilio) processing discards","year":2023,"lang":"en","type":"article","venue":"Memorial University Research Repository (Memorial University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Memorial University of Newfoundland","funders":"","keywords":"Discards; Aquaculture; Biorefinery; Snow; Hazardous waste; Agriculture; Shellfish; Bioproducts","score_opus":0.029889009537966617,"score_gpt":0.27354753668701615,"score_spread":0.24365852714904954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948047220","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9804764,0.00045963263,0.002287713,0.0006830383,0.000018294058,0.00061989255,0.0001515939,0.000028116843,0.015275292],"genre_scores_gemma":[0.9712225,0.0014556746,0.0140609285,0.0005293516,0.0000066301404,0.00035508984,0.00020624051,0.000011064444,0.012152506],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9989409,0.000308384,0.000041920404,0.00010606606,0.0003165237,0.0002862716],"domain_scores_gemma":[0.9992526,0.00013938107,0.0001645791,0.00003702986,0.0002830436,0.00012332154],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018925817,0.00051214633,0.00017715749,0.00075138867,0.0013209443,0.0017789903,0.00061710563,0.0005091422,0.0015338594],"category_scores_gemma":[0.0008578778,0.00016846278,0.00050048257,0.00038964732,0.0012206748,0.0006013302,0.0014847366,0.00056986255,0.00019574356],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0017114546,0.007036195,0.09056588,0.0025238937,0.00021998522,0.0026643868,0.014388628,0.013954705,0.57179326,0.005377948,0.004981911,0.28478184],"study_design_scores_gemma":[0.0003727712,0.025806313,0.3793165,0.0012337642,0.0005096418,0.0009109717,0.059864774,0.009103499,0.39740053,0.0013497481,0.12382314,0.00030833666],"about_ca_topic_score_codex":0.18502524,"about_ca_topic_score_gemma":0.481713,"teacher_disagreement_score":0.8149748,"about_ca_system_score_codex":0.010752242,"about_ca_system_score_gemma":0.0076657296,"threshold_uncertainty_score":0.36789656},"labels":[],"label_agreement":null},{"id":"W6948123523","doi":"10.48550/arxiv.1102.3694","title":"On the evolution of the molecular gas fraction of star forming galaxies","year":2011,"lang":"en","type":"preprint","venue":"arXiv (Cornell University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Galaxy; Star formation; Plateau de Bure Interferometer; Galaxy formation and evolution; Halo; Peculiar galaxy; Elliptical galaxy","score_opus":0.04309692735787385,"score_gpt":0.18724408945086643,"score_spread":0.14414716209299258,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948123523","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9985839,0.00014220361,0.00034674353,0.000038857514,5.2095714e-7,0.0000011777092,0.0005396051,0.000022461782,0.0003244477],"genre_scores_gemma":[0.99887437,0.00006324851,0.0002953707,0.000005506937,0.0000028536995,0.00000148224,0.0006650235,0.0000035232567,0.00008868375],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99984026,0.00002603189,0.0000069598973,0.00006575572,0.00002502985,0.000035986348],"domain_scores_gemma":[0.9975922,0.0012185067,0.000672949,0.0001667491,0.0001898061,0.0001597264],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004095292,0.00016611876,0.00019850495,0.0019573877,0.00023739027,0.00081540167,0.00026005052,0.00029410655,0.00065736665],"category_scores_gemma":[0.003161101,0.00011058271,0.00013510171,0.0013439097,0.0003536024,0.00049742893,0.00033644194,0.00015491337,0.00016253388],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011082074,0.000017204462,0.9748824,0.000016849894,0.000047461104,0.000107225904,0.00024305891,0.002921827,0.009636932,0.00070775376,0.00015538733,0.011153165],"study_design_scores_gemma":[0.0000019723905,0.000021196262,0.9884803,0.0000024916715,0.000012484988,0.00021288487,0.000085166,0.008406862,0.0018911444,0.00045396903,0.00042610487,0.000005265801],"about_ca_topic_score_codex":0.0039335364,"about_ca_topic_score_gemma":0.0038328548,"teacher_disagreement_score":0.0039335364,"about_ca_system_score_codex":0.00044499183,"about_ca_system_score_gemma":0.00009760309,"threshold_uncertainty_score":0.0078213215},"labels":[],"label_agreement":null},{"id":"W6948188175","doi":"10.4224/21272474","title":"Low cost high performance cluster computing: capabilities, techniques and application to product performance modeling","year":2002,"lang":"en","type":"report","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"National Research Council Canada","funders":"","keywords":"Cluster (spacecraft); Software; Set (abstract data type); Product (mathematics); Key (lock)","score_opus":0.02001915514728021,"score_gpt":0.27113549495825223,"score_spread":0.25111633981097203,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948188175","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.007902919,0.0013314704,0.9732617,0.00083478144,0.000079814665,0.000108187356,0.0007042297,0.0028399639,0.012936878],"genre_scores_gemma":[0.2983184,0.004619776,0.6805893,0.000121266596,0.00018693932,0.0004896234,0.0023295814,0.0016206392,0.011724432],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.999097,0.00022822972,0.000033160562,0.00008312708,0.00050656276,0.00005206754],"domain_scores_gemma":[0.9989642,0.00047173255,0.000075654985,0.00020939653,0.00023791383,0.000041102747],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010017481,0.0008740799,0.0005951325,0.001806019,0.00096861436,0.0021029436,0.0021960824,0.0007467744,0.004837272],"category_scores_gemma":[0.0033688485,0.0005711684,0.0006621308,0.006565366,0.00063066586,0.0019453613,0.0008178864,0.0011450001,0.0021103777],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000076986034,0.00006413759,0.003030456,0.00037142885,0.000073197225,0.00012910021,0.00017190255,0.63644433,0.0027469231,0.14918932,0.024316497,0.18338577],"study_design_scores_gemma":[0.000007515252,0.000015315929,0.0009084952,0.000018815088,0.0000112231155,0.000056871453,0.000032598102,0.9455289,0.0011091812,0.034826424,0.017463902,0.000020742995],"about_ca_topic_score_codex":0.020174496,"about_ca_topic_score_gemma":0.015171526,"teacher_disagreement_score":0.020174496,"about_ca_system_score_codex":0.0019023011,"about_ca_system_score_gemma":0.0018897966,"threshold_uncertainty_score":0.040114164},"labels":[],"label_agreement":null},{"id":"W6948297781","doi":"10.5064/f6z31wj1/eaanir","title":"Johnson_House41_JUSTHR_038ev.pdf","year":2023,"lang":"fr","type":"dataset","venue":"Syracuse University Qualitative Data Repository","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Process (computing); Identification (biology); Product (mathematics)","score_opus":0.07050987181439942,"score_gpt":0.34553004200518345,"score_spread":0.275020170190784,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948297781","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00029655008,0.000023518929,0.00011872851,0.00010538312,0.000014993576,0.000014345141,0.99656457,0.00081063237,0.0020514193],"genre_scores_gemma":[0.0007590099,0.000036302812,0.00037591008,0.0000458907,0.000005108887,0.000057572463,0.9959858,0.00011292886,0.002621465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99918574,0.00011787506,0.0000833422,0.0002018907,0.00025462566,0.00015649227],"domain_scores_gemma":[0.9977055,0.0006209192,0.00015619534,0.00058784615,0.0006325055,0.00029711562],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0011937477,0.0012928909,0.00069079455,0.0035623466,0.00092147663,0.0019992618,0.0017704185,0.0011212426,0.13945395],"category_scores_gemma":[0.0050942665,0.0006250985,0.00071645156,0.0049456833,0.00041986958,0.0011608099,0.0015642317,0.0009837933,0.12852392],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000047787045,0.000018748275,0.0005185335,0.00017269311,0.000009357182,0.000013123967,0.000020839318,0.00012493721,0.00013480749,0.00045500215,0.9949719,0.0035122815],"study_design_scores_gemma":[0.00019951875,0.000023314795,0.0050277547,0.00012260863,0.000017726847,0.00005836773,0.00015752525,0.0007401612,0.0009871683,0.00081252557,0.99182206,0.00003116228],"about_ca_topic_score_codex":0.05850687,"about_ca_topic_score_gemma":0.089461595,"teacher_disagreement_score":0.86054605,"about_ca_system_score_codex":0.001976771,"about_ca_system_score_gemma":0.0022578954,"threshold_uncertainty_score":0.46651995},"labels":[],"label_agreement":null},{"id":"W6948306908","doi":"10.5061/dryad.05qfttf0n","title":"Fish mock community with 41 species from 13 orders","year":2020,"lang":"en","type":"dataset","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University; University of Guelph; Shared Hierarchical Academic Research Computing Network","funders":"","keywords":"Fish <Actinopterygii>; DNA sequencing; Replicate; Fishing; Diversity (politics); Muscle tissue","score_opus":0.04364010703638297,"score_gpt":0.24769246881035814,"score_spread":0.20405236177397518,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948306908","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061659394,0.00024314946,0.00031033103,0.00005794632,0.000025312796,0.000042701475,0.99183613,0.00029564786,0.0010227617],"genre_scores_gemma":[0.0019338136,0.000055989312,0.00054492784,0.000028795988,0.0000023394625,0.00008711548,0.9967687,0.000028756911,0.0005494789],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991334,0.00010114256,0.00008745265,0.00034392567,0.00019116634,0.00014290115],"domain_scores_gemma":[0.9993043,0.00012788513,0.00007163492,0.00017158645,0.00021156156,0.0001130131],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010826734,0.0018000398,0.0014120796,0.0033043392,0.0012151593,0.0009901744,0.0020724493,0.0014324009,0.01987572],"category_scores_gemma":[0.0019082034,0.00064155,0.0012558498,0.004287868,0.00062377524,0.00080134295,0.0020751266,0.0011192369,0.023582736],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011212763,0.00022338248,0.033369705,0.0043383925,0.0005029577,0.0006594674,0.0004933249,0.002665929,0.006581813,0.0014120307,0.9143292,0.03430254],"study_design_scores_gemma":[0.00034432093,0.00007307949,0.0607332,0.0004977492,0.00019505435,0.00052064616,0.00058302464,0.0012098807,0.002068319,0.0010759825,0.93261135,0.000087372915],"about_ca_topic_score_codex":0.052377805,"about_ca_topic_score_gemma":0.121887825,"teacher_disagreement_score":0.052377805,"about_ca_system_score_codex":0.0016313663,"about_ca_system_score_gemma":0.0019984029,"threshold_uncertainty_score":0.104145885},"labels":[],"label_agreement":null},{"id":"W6948392041","doi":"10.5061/dryad.j0t179b/2","title":"seedweight","year":2018,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Raw material; Raw data; Product (mathematics); Seed testing","score_opus":0.022351944630091852,"score_gpt":0.27122296463996504,"score_spread":0.24887102000987318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948392041","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001533125,0.00032698177,0.015179392,0.00042419357,0.00035267184,0.0002902152,0.8146825,0.12282719,0.044383664],"genre_scores_gemma":[0.005768353,0.00027654815,0.02189744,0.00044754616,0.000056603985,0.00047091173,0.89764595,0.036474146,0.03696253],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993698,0.00005529179,0.000058498317,0.00019654776,0.00025539027,0.00006441791],"domain_scores_gemma":[0.99860114,0.0003482921,0.000079100166,0.00038140683,0.00048628417,0.00010381787],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00088262407,0.0019022984,0.0008554942,0.002410784,0.00089855853,0.0023366106,0.0021220774,0.001038831,0.33829233],"category_scores_gemma":[0.0053861025,0.0009072376,0.0012970257,0.002438588,0.0002383705,0.002929406,0.0026979297,0.001144487,0.3109368],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014439822,0.000037409587,0.00088423677,0.00049967336,0.00003191397,0.000054542965,0.00007008888,0.00027790942,0.0011246018,0.0023965102,0.9620336,0.03244513],"study_design_scores_gemma":[0.00006330229,0.000019691017,0.0015684421,0.00007237537,0.000018450693,0.00010430802,0.00004619759,0.0010191122,0.0020150365,0.0032130638,0.9918331,0.000026879517],"about_ca_topic_score_codex":0.0032061536,"about_ca_topic_score_gemma":0.0061742766,"teacher_disagreement_score":0.66170764,"about_ca_system_score_codex":0.0008872637,"about_ca_system_score_gemma":0.0010854133,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W6948411440","doi":"","title":"Enteroaggregative Escherichia coli as a Major Etiologic Agent in Traveler's Diarrhea in 3 Regions of the World","year":2017,"lang":"en","type":"other","venue":"reroDoc Digital Library","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Enteroaggregative Escherichia coli; Diarrhea; Traveler's diarrhea; Pathogen; Escherichia coli","score_opus":0.013888957140308697,"score_gpt":0.24805626266591158,"score_spread":0.2341673055256029,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948411440","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99690324,0.0008530392,0.00014959821,0.00011957908,0.0000021167264,0.000014261215,0.0010256929,0.00001396877,0.0009184011],"genre_scores_gemma":[0.9967526,0.000910953,0.0004823976,0.000055807155,0.0000053358694,0.000011859195,0.0013957659,0.000006160701,0.0003789843],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9997353,0.000055747638,0.00003623108,0.000055173346,0.000064086365,0.000053455144],"domain_scores_gemma":[0.99956316,0.000052576404,0.0002505796,0.000027046675,0.000056717487,0.00004985564],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00034677293,0.00017549339,0.00016850166,0.0022635143,0.00038625998,0.0005629468,0.0001588148,0.00016531827,0.0009834011],"category_scores_gemma":[0.0009607459,0.000092028546,0.00019183493,0.0038631451,0.00031087387,0.00026473336,0.00067903573,0.00015512727,0.0001537627],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018320793,0.000039856306,0.9724832,0.00012278094,0.00003488378,0.00047559783,0.0009353402,0.00007307843,0.0020253698,0.00006829209,0.00044125202,0.023117235],"study_design_scores_gemma":[0.0000032551,0.000015784928,0.99699545,0.000025894675,0.000034771332,0.00048097075,0.0010153584,0.00009521527,0.00026183767,0.000037112226,0.0010306613,0.0000037155942],"about_ca_topic_score_codex":0.033246364,"about_ca_topic_score_gemma":0.043031126,"teacher_disagreement_score":0.033246364,"about_ca_system_score_codex":0.00038073308,"about_ca_system_score_gemma":0.0006991759,"threshold_uncertainty_score":0.06610572},"labels":[],"label_agreement":null},{"id":"W6948430855","doi":"10.5066/p9w03vwz","title":"Whole rock geochemistry data from the Ordovician Bronson Hill arc and Silurian and Devonian Connecticut Valley - Gasp&amp;amp;amp;eacute; trough, Vermont and New Hampshire","year":2023,"lang":"en","type":"dataset","venue":"USGS DOI Tool Production Environment","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ordovician; Devonian; Island arc; Carbonate rock; Paleozoic; Precambrian","score_opus":0.04567319844691532,"score_gpt":0.26899848681947336,"score_spread":0.22332528837255805,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948430855","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00026838682,0.00003394587,0.000037747304,0.000037870435,0.000010041292,0.0000070677406,0.9988539,0.0001054164,0.0006455194],"genre_scores_gemma":[0.0004163101,0.000030348843,0.00015095124,0.000018220086,0.0000020691484,0.000029383931,0.99881613,0.000025580697,0.0005109533],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992538,0.000072108676,0.00007742354,0.0002260317,0.00024237871,0.00012820518],"domain_scores_gemma":[0.9985464,0.00028218539,0.00016727766,0.0002675803,0.0005192045,0.00021733514],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008507327,0.0017371766,0.0009789248,0.0041732374,0.0009815028,0.0016233681,0.002398508,0.0009963043,0.03403747],"category_scores_gemma":[0.0031785183,0.0006268262,0.0006090347,0.009949501,0.00051185983,0.0006630898,0.0015981696,0.0012897392,0.038197257],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054569085,0.000020548163,0.0020654502,0.00043458355,0.000025303925,0.000055925204,0.000059677815,0.00032532727,0.00023539148,0.0005333942,0.9937928,0.002397067],"study_design_scores_gemma":[0.00012883132,0.000009324218,0.011239768,0.00024378207,0.000024143515,0.000078935955,0.00021026374,0.0002889941,0.0006808018,0.0006018567,0.98647004,0.000023241546],"about_ca_topic_score_codex":0.12352508,"about_ca_topic_score_gemma":0.23464338,"teacher_disagreement_score":0.12352508,"about_ca_system_score_codex":0.0027400295,"about_ca_system_score_gemma":0.004610242,"threshold_uncertainty_score":0.24561214},"labels":[],"label_agreement":null},{"id":"W6948452307","doi":"10.5281/zenodo.11525094","title":"What are the main ingredients in Manup Male Enhancement Gummies?","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Testosterone (patch); Lift (data mining); Performance enhancement; Order (exchange); Human enhancement","score_opus":0.02728643448568354,"score_gpt":0.2681130781454389,"score_spread":0.24082664365975534,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948452307","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.029952344,0.21968304,0.020330418,0.0337963,0.015376787,0.0028721797,0.010646243,0.013262859,0.65408],"genre_scores_gemma":[0.062191274,0.1264335,0.030520635,0.02056103,0.004546914,0.0012699136,0.004416758,0.0032836804,0.74677634],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996032,0.000064835585,0.0000324596,0.00004617683,0.00021650415,0.00003688155],"domain_scores_gemma":[0.9994023,0.00010963051,0.000099904,0.000034723955,0.0002275975,0.00012588716],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0008083443,0.0010299477,0.0008218466,0.0014783117,0.0012342563,0.0021019543,0.000834531,0.0022455638,0.20103987],"category_scores_gemma":[0.0016328744,0.0005242562,0.00076698244,0.00071405666,0.0005106104,0.003440508,0.0009577852,0.0022856363,0.10106529],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007590245,0.0006640938,0.00086911285,0.0066814832,0.00004765846,0.00059451454,0.0007343821,0.000084540756,0.014974887,0.005325244,0.34818268,0.62108237],"study_design_scores_gemma":[0.00003815164,0.0001895031,0.0011720742,0.0006460878,0.000040559684,0.00034768818,0.00030930826,0.000046785844,0.0035620842,0.0005962275,0.9930253,0.000026178337],"about_ca_topic_score_codex":0.0023317812,"about_ca_topic_score_gemma":0.005753732,"teacher_disagreement_score":0.79896015,"about_ca_system_score_codex":0.00039605927,"about_ca_system_score_gemma":0.00074729795,"threshold_uncertainty_score":0.6725453},"labels":[],"label_agreement":null},{"id":"W6948468743","doi":"10.5061/dryad.3j9kd51sh","title":"Mass spectrometry of axonemes from Tetrahymena thermophila CU428 and acetylation mutants","year":2024,"lang":"en","type":"dataset","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Canadian Institutes of Health Research","keywords":"Acetylation; Tetrahymena; Microtubule; Tubulin; Mutant; Cilium; Lysine","score_opus":0.025286646458487737,"score_gpt":0.31076985952284086,"score_spread":0.28548321306435315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948468743","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.10716956,0.0018934121,0.0010285235,0.00023410014,0.00007623639,0.00005118123,0.88599885,0.0017650292,0.0017831179],"genre_scores_gemma":[0.0264527,0.00036236166,0.0025193144,0.00006165024,0.000006725771,0.00008541445,0.9698242,0.000055754375,0.0006318669],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9995028,0.00004108982,0.000070848044,0.00020531494,0.000118563315,0.00006135481],"domain_scores_gemma":[0.99947864,0.00011724423,0.00010043362,0.000103222264,0.00014465571,0.00005589301],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00057483546,0.0013115198,0.0008879598,0.002741402,0.0011137022,0.00091576943,0.0009093652,0.0012704537,0.001946595],"category_scores_gemma":[0.001214465,0.00025580465,0.001041129,0.0034054564,0.00034582295,0.00047374546,0.0009782377,0.00069814385,0.002150935],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.006993647,0.0009120061,0.10242012,0.015145621,0.0022141973,0.006773467,0.0011955196,0.011128472,0.29085055,0.003567134,0.45594463,0.10285459],"study_design_scores_gemma":[0.0005701975,0.00040190955,0.36683828,0.000588888,0.0009138969,0.002893485,0.0010641513,0.014987075,0.06913875,0.002811349,0.5395766,0.00021551529],"about_ca_topic_score_codex":0.01275377,"about_ca_topic_score_gemma":0.020122478,"teacher_disagreement_score":0.01275377,"about_ca_system_score_codex":0.00093234726,"about_ca_system_score_gemma":0.0011330481,"threshold_uncertainty_score":0.025359094},"labels":[],"label_agreement":null},{"id":"W6948542954","doi":"10.5281/zenodo.10123357","title":"Fig. 39 in A revision of Scipopus Enderlein including the subgenera Scipopus s. str., Phaeopterina Frey and Parascipopus subgen. nov. (Diptera, Micropezidae, Taeniapterinae)","year":2023,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph","funders":"","keywords":"Spermatheca; Subgenus; Dorsum; Male genitalia; Sex organ; Head (geology)","score_opus":0.04407853115487176,"score_gpt":0.2839744164672181,"score_spread":0.23989588531234635,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948542954","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.036943875,0.011088456,0.022324221,0.0017029351,0.005975479,0.0010255034,0.053667914,0.0035467176,0.8637248],"genre_scores_gemma":[0.2766488,0.021220496,0.08481302,0.001893577,0.0030422555,0.0013944671,0.16296883,0.0019847027,0.44603378],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9998388,0.000023658506,0.000018902037,0.00003746407,0.000060362105,0.00002077577],"domain_scores_gemma":[0.999821,0.00002638188,0.000031254116,0.00002745302,0.00007254516,0.000021376438],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00024128292,0.00086202234,0.00026668425,0.0038938425,0.0012446216,0.0008714222,0.0006706569,0.00038402455,0.047192235],"category_scores_gemma":[0.00061506225,0.0001793509,0.00033581723,0.002741286,0.0010425613,0.0010137646,0.0007076508,0.00089114666,0.017259292],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024068258,0.000052426993,0.009422022,0.0017142355,0.000061824954,0.00074322696,0.0014077356,0.000779732,0.005542786,0.02353492,0.33050936,0.625991],"study_design_scores_gemma":[0.0000055598734,0.00001657113,0.0109411,0.00014544756,0.000019046118,0.00051953027,0.00018941339,0.0001407978,0.00028699395,0.0006118475,0.98711604,0.0000075382923],"about_ca_topic_score_codex":0.011372067,"about_ca_topic_score_gemma":0.021884853,"teacher_disagreement_score":0.047192235,"about_ca_system_score_codex":0.0007582616,"about_ca_system_score_gemma":0.000984739,"threshold_uncertainty_score":0.15787381},"labels":[],"label_agreement":null},{"id":"W6948633003","doi":"10.5255/ukda-sn-6702-35","title":"Monthly Wages and Salaries Survey, 2000-2023: Secure Access","year":2023,"lang":"en","type":"dataset","venue":"UK Data Archive","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Earnings; Index (typography); Stratified sampling; Quarter (Canadian coin); Cover (algebra); Business information; Small business","score_opus":0.06211337689932499,"score_gpt":0.34257898814212856,"score_spread":0.28046561124280356,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948633003","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003787083,0.000078419835,0.00015530517,0.0001580976,0.000040339284,0.00021313142,0.98752844,0.00014169006,0.007897451],"genre_scores_gemma":[0.009464524,0.00021145685,0.0006048654,0.00025905072,0.00005188125,0.0010125382,0.97178054,0.00007041517,0.016544664],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99862087,0.00017156561,0.00025729687,0.00019790934,0.00054200756,0.00021042419],"domain_scores_gemma":[0.9956173,0.00034231832,0.00079136284,0.000274129,0.0026047782,0.00037010538],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0016728005,0.0008161923,0.0005808187,0.003469444,0.00042332022,0.0011104039,0.0012486243,0.0006939757,0.035773087],"category_scores_gemma":[0.0057582865,0.0004924031,0.0003682628,0.010011562,0.00015607341,0.001201715,0.0010713837,0.0010640058,0.04969026],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025532907,0.00017368428,0.033241022,0.00041214368,0.00003110257,0.00004072626,0.00012442155,0.000297608,0.00016857062,0.0004870636,0.9443926,0.020375652],"study_design_scores_gemma":[0.0002788017,0.000119678625,0.5450493,0.00026795635,0.000028253115,0.00006838848,0.0003773968,0.00079408963,0.00029744365,0.00043197465,0.45224866,0.00003808956],"about_ca_topic_score_codex":0.07582137,"about_ca_topic_score_gemma":0.08166224,"teacher_disagreement_score":0.07582137,"about_ca_system_score_codex":0.0017338127,"about_ca_system_score_gemma":0.0022882202,"threshold_uncertainty_score":0.15076005},"labels":[],"label_agreement":null},{"id":"W6948676687","doi":"10.5281/zenodo.11988003","title":"road access binder quebec pdf","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Work (physics); Limiting; Filter (signal processing); Frame (networking)","score_opus":0.03415362984051714,"score_gpt":0.28869357704121873,"score_spread":0.2545399472007016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948676687","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002155275,0.00026395498,0.00018717568,0.00046932176,0.00052634074,0.000086760425,0.0040152348,0.0012604765,0.9929751],"genre_scores_gemma":[0.00071951735,0.00011643138,0.0000995018,0.00016925453,0.0000347663,0.000014121578,0.00093277334,0.00038718333,0.99752647],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992125,0.000035866102,0.000018437635,0.00010831989,0.00045546566,0.00016935721],"domain_scores_gemma":[0.9967445,0.00014495644,0.000045728117,0.00023330137,0.0022966557,0.0005347482],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0005729721,0.001104465,0.001110277,0.0020307272,0.004668232,0.0069639413,0.002001378,0.0025176613,0.9431234],"category_scores_gemma":[0.0027835313,0.0007195203,0.00079267466,0.0028227267,0.00087273173,0.0028475446,0.001925927,0.0020080856,0.84503156],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018589411,0.000018581948,0.00006824938,0.000047179212,0.0000012823623,0.000027011267,0.000022235778,0.000031009236,0.00014019017,0.00097236765,0.9801517,0.018501569],"study_design_scores_gemma":[0.0000061125493,0.0000049773835,0.00031736662,0.000033781194,0.0000011567781,0.00001757057,0.000049318318,0.000027231696,0.000097168326,0.00010328797,0.99933594,0.000006078345],"about_ca_topic_score_codex":0.47117198,"about_ca_topic_score_gemma":0.7246548,"teacher_disagreement_score":0.47117198,"about_ca_system_score_codex":0.010178365,"about_ca_system_score_gemma":0.008842498,"threshold_uncertainty_score":0.9368589},"labels":[],"label_agreement":null},{"id":"W6948681032","doi":"10.5061/dryad.dm3hb40","title":"Data from: Genome-wide assessment of diversity and divergence among extant Galápagos giant tortoise species","year":2018,"lang":"en","type":"dataset","venue":"Data Archiving and Networked Services (DANS)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Tortoise; Genetic diversity; Microsatellite; Mitochondrial DNA; Conservation genetics; Population; Population genetics; Extant taxon","score_opus":0.04649933439612716,"score_gpt":0.29074562578088253,"score_spread":0.24424629138475537,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948681032","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.116977245,0.00030574703,0.0036932293,0.00044454166,0.0000863604,0.00010783937,0.86891204,0.0014134728,0.008059545],"genre_scores_gemma":[0.13146195,0.00039598497,0.021811629,0.0002068462,0.000058307494,0.00041550468,0.84206426,0.00035512896,0.003230481],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.99946517,0.00006476706,0.00007386402,0.00017534434,0.0001677597,0.000053145057],"domain_scores_gemma":[0.9983152,0.0003236892,0.00033399288,0.0003748188,0.0004359812,0.00021637455],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0009825858,0.0005248329,0.00046563323,0.00414828,0.0005532804,0.0010749383,0.00080487167,0.00062058197,0.013976593],"category_scores_gemma":[0.0027638844,0.0003019647,0.0003942799,0.0048665027,0.00030941644,0.0005363099,0.0011948865,0.00066095655,0.0036085339],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0016490101,0.0003501544,0.27369234,0.004922899,0.0007564998,0.0014799094,0.004860578,0.0038071137,0.18161066,0.003647875,0.3234526,0.19977036],"study_design_scores_gemma":[0.00014004229,0.000079687845,0.8149829,0.0003403906,0.00020474521,0.0007114839,0.00083190814,0.001939331,0.011494305,0.0019593819,0.16721392,0.00010181855],"about_ca_topic_score_codex":0.008759312,"about_ca_topic_score_gemma":0.014681743,"teacher_disagreement_score":0.013976593,"about_ca_system_score_codex":0.00048204104,"about_ca_system_score_gemma":0.00075425615,"threshold_uncertainty_score":0.046756387},"labels":[],"label_agreement":null},{"id":"W6948695983","doi":"10.5281/zenodo.11924751","title":"sheila heti maternidad pdf","year":2024,"lang":"pl","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Buckle; Limiting; Ceylon; Pretext; Work (physics)","score_opus":0.023750935617780758,"score_gpt":0.25835682649856984,"score_spread":0.2346058908807891,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948695983","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00023822184,0.00052686257,0.00025175017,0.0016692725,0.0014844269,0.00007202288,0.0026034622,0.0020634308,0.99109054],"genre_scores_gemma":[0.0010125452,0.0005601849,0.00028717148,0.0006429445,0.00024987658,0.000031145235,0.0012357397,0.0007497517,0.9952306],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99952364,0.000042305215,0.000017125387,0.000055542478,0.00024284379,0.00011851333],"domain_scores_gemma":[0.9978923,0.00020697973,0.00008718233,0.00015429838,0.00076307665,0.0008960659],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00056663505,0.0008601086,0.0007094094,0.0015316432,0.0017214073,0.0050027855,0.0013159072,0.0015215831,0.927036],"category_scores_gemma":[0.003438924,0.00060335075,0.00058763986,0.0013207833,0.00055327435,0.0026428492,0.0038650616,0.0023685226,0.84773153],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000024361316,0.00002113615,0.00006571041,0.00007826184,9.051887e-7,0.00007042702,0.000047313006,0.0000125553515,0.00010723467,0.000863528,0.96034217,0.038366538],"study_design_scores_gemma":[0.00000513246,0.000007084369,0.00019640283,0.000057189307,6.752326e-7,0.00006895618,0.0000694933,0.000007062944,0.00005367655,0.00011720112,0.99941325,0.0000038826092],"about_ca_topic_score_codex":0.0059844474,"about_ca_topic_score_gemma":0.015515093,"teacher_disagreement_score":0.07296401,"about_ca_system_score_codex":0.0015381243,"about_ca_system_score_gemma":0.0021848122,"threshold_uncertainty_score":0.1040743},"labels":[],"label_agreement":null},{"id":"W6948703995","doi":"10.5281/zenodo.11836462","title":"+&gt;![[.WATCH.] Breathe (2024) (FULLMOVIE) ONLINE ON 123MOVIES","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Staring; Limiting; Pretext; Circumstantial evidence; Ridiculous; Paraphernalia","score_opus":0.029507158072553927,"score_gpt":0.27308302936810047,"score_spread":0.24357587129554653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948703995","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0004029814,0.00027408142,0.000965525,0.00071800256,0.0012831986,0.00011602305,0.0056118146,0.009705097,0.9809232],"genre_scores_gemma":[0.0016410165,0.00019337235,0.0005331455,0.00057242997,0.00021389044,0.00006189452,0.0032721164,0.005432924,0.98807913],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997596,0.000024495437,0.000008573765,0.000046477522,0.00009911764,0.00006167892],"domain_scores_gemma":[0.99885285,0.000082175044,0.000032985627,0.00013319324,0.00040742985,0.00049137545],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0003399573,0.00083895813,0.0005309304,0.0009997974,0.0016504225,0.004794744,0.0010240824,0.0015311897,0.9496688],"category_scores_gemma":[0.0021259442,0.0004951174,0.00064593944,0.00080985273,0.00035287804,0.004413697,0.0037844526,0.0013232569,0.9352007],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000017863198,0.000009825082,0.00003534138,0.00004682156,0.000001059586,0.000016793008,0.00003257158,0.000008967348,0.00016471666,0.0005077498,0.98251086,0.01664745],"study_design_scores_gemma":[0.0000070624137,0.000008254719,0.00018613433,0.00003670334,0.0000013393934,0.000027027492,0.000058541733,0.000025986856,0.000113839815,0.00020705877,0.99932206,0.000006112027],"about_ca_topic_score_codex":0.0033352994,"about_ca_topic_score_gemma":0.00837233,"teacher_disagreement_score":0.050331175,"about_ca_system_score_codex":0.0006157359,"about_ca_system_score_gemma":0.00039240636,"threshold_uncertainty_score":0.07179129},"labels":[],"label_agreement":null},{"id":"W6948718858","doi":"10.5281/zenodo.10021247","title":"Jaeyun-Song/RGE: Repulsive-Graph-Rectification: Official code","year":2023,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Kootenay Association for Science & Technology","funders":"","keywords":"Code (set theory); Graph; Node (physics); Rectification; Context (archaeology)","score_opus":0.03680980500654625,"score_gpt":0.276402430186708,"score_spread":0.23959262518016178,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948718858","genre_codex":"software","genre_gemma":"software","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"software","genre_consensus":"software","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013136545,0.0002693514,0.4337071,0.0008074699,0.000771891,0.00035338817,0.039276306,0.50357985,0.019920912],"genre_scores_gemma":[0.024922978,0.0005700055,0.46741524,0.0010279536,0.0002867075,0.001115641,0.16529375,0.28141284,0.057954848],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9981382,0.00021971657,0.00025523469,0.00036539615,0.00083056686,0.00019096627],"domain_scores_gemma":[0.9977481,0.00065346353,0.00012560356,0.0006172155,0.0007048534,0.00015082578],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0018453917,0.0022100569,0.00091500196,0.00225697,0.00090071734,0.0029361772,0.0028771604,0.0022041448,0.12883869],"category_scores_gemma":[0.008116735,0.0016411031,0.0020039377,0.0018157674,0.00091141113,0.003640704,0.003111294,0.0031903926,0.120494105],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00030398046,0.00007589763,0.0005812853,0.0007942145,0.00004636999,0.0001505862,0.00021484907,0.0032074605,0.004537671,0.033742756,0.8349538,0.12139111],"study_design_scores_gemma":[0.00021480156,0.000040070765,0.0006927442,0.00016972735,0.00002877784,0.00031798246,0.000055638324,0.033994988,0.012918729,0.036028072,0.91542375,0.00011470469],"about_ca_topic_score_codex":0.005221572,"about_ca_topic_score_gemma":0.0067188255,"teacher_disagreement_score":0.12883869,"about_ca_system_score_codex":0.0012455413,"about_ca_system_score_gemma":0.0020642302,"threshold_uncertainty_score":0.43100834},"labels":[],"label_agreement":null},{"id":"W6948747039","doi":"10.5281/zenodo.12274932","title":"pmi salary survey pdf","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Salary; Leverage (statistics); Survey data collection; Human resource management; Wages and salaries","score_opus":0.03491278904603738,"score_gpt":0.2694294908420082,"score_spread":0.23451670179597078,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948747039","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006627582,0.00029905877,0.00051057787,0.003386956,0.0018757468,0.00037339373,0.016798388,0.0023677784,0.9737254],"genre_scores_gemma":[0.0022608724,0.00035020366,0.00049575645,0.0013578313,0.0005427595,0.0002799866,0.0077148587,0.0006815943,0.98631614],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99823976,0.00014291487,0.000054727137,0.00015980616,0.0011646026,0.00023814863],"domain_scores_gemma":[0.9943336,0.00039681385,0.00017446536,0.00033545698,0.003254736,0.0015048871],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0010844402,0.00076760066,0.00076509936,0.0019349921,0.002012041,0.0039013936,0.0012725651,0.0013381431,0.8560783],"category_scores_gemma":[0.007170066,0.0005832996,0.0004576537,0.0029365318,0.00030465744,0.0024020635,0.0023705843,0.0025311909,0.8488],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000060353805,0.000012938364,0.000056357232,0.000019570303,3.112986e-7,0.0000026666282,0.000006378336,0.0000050264434,0.000025206029,0.00015781632,0.9901518,0.009555982],"study_design_scores_gemma":[0.000007305216,0.000017902023,0.0011428475,0.00002850392,8.3921617e-7,0.000016518357,0.0000616425,0.000036360892,0.000055362107,0.00009776592,0.99853015,0.000004706178],"about_ca_topic_score_codex":0.00644666,"about_ca_topic_score_gemma":0.010370784,"teacher_disagreement_score":0.14392167,"about_ca_system_score_codex":0.0015950211,"about_ca_system_score_gemma":0.0030027796,"threshold_uncertainty_score":0.2052868},"labels":[],"label_agreement":null},{"id":"W6948831981","doi":"10.5281/zenodo.10073499","title":"ENSINO INTERCULTURAL DE ESTEREÓTIPOS EM ATIVIDADES DE LÍNGUA ESTRANGEIRA: UMA REFLEXÃO CRÍTICA E DIALÓGICA PARA UMA POSIÇÃO RESPONSIVA EM SALA DE AULA","year":2023,"lang":"pt","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Context (archaeology); Order (exchange); Happening","score_opus":0.06289752165333982,"score_gpt":0.3191198383502249,"score_spread":0.25622231669688506,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948831981","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.46758273,0.007655075,0.01582367,0.05424018,0.0007586485,0.0000819317,0.0001882982,0.00014825123,0.45352128],"genre_scores_gemma":[0.9734851,0.0011358622,0.00067907077,0.0011472485,0.00008496281,0.000026618392,0.000019617162,0.000071127055,0.023350308],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9933776,0.0042675543,0.00011221735,0.0006150215,0.000952562,0.00067499105],"domain_scores_gemma":[0.99045205,0.0053899204,0.001052738,0.000685819,0.0016343829,0.0007851362],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007940874,0.00066798803,0.00053076266,0.0024885943,0.020112947,0.016931716,0.0014962943,0.0030148374,0.014315395],"category_scores_gemma":[0.010784212,0.0004536882,0.00033780435,0.0025296214,0.03043931,0.011392891,0.009540428,0.005317461,0.0011355376],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000032980508,0.000018114362,0.003010547,0.00009453693,0.000005696529,0.00049060857,0.864855,0.000059802736,0.000787803,0.119305,0.0029501605,0.008389815],"study_design_scores_gemma":[0.0000075744288,0.000023564335,0.0057116267,0.0003244642,0.000019327306,0.0003825495,0.76641583,0.00030884668,0.0009816077,0.019327966,0.2064592,0.00003750317],"about_ca_topic_score_codex":0.094752945,"about_ca_topic_score_gemma":0.15162158,"teacher_disagreement_score":0.094752945,"about_ca_system_score_codex":0.02041366,"about_ca_system_score_gemma":0.010245735,"threshold_uncertainty_score":0.18840283},"labels":[],"label_agreement":null},{"id":"W6948922682","doi":"10.5281/zenodo.10667955","title":"Acalypha cardielii I. Montero & G. A. Levin","year":2023,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Canadian Museum of Nature","funders":"","keywords":"IUCN Red List; Protected area; Deciduous; Endemism; National park; IUCN protected area categories; Habitat; Deforestation (computer science); Habitat destruction","score_opus":0.04289258725603524,"score_gpt":0.27563792188726255,"score_spread":0.23274533463122732,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6948922682","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.13443616,0.081894085,0.002521781,0.0047346116,0.0016479911,0.0005425974,0.0015669208,0.00023537558,0.7724205],"genre_scores_gemma":[0.84872466,0.034294337,0.008352197,0.0028731113,0.0015824953,0.00026195685,0.001910646,0.000049860024,0.10195066],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.999895,0.000013653021,0.000005786579,0.000029326317,0.000033019067,0.000023175697],"domain_scores_gemma":[0.9999262,0.000017048053,0.000024242418,0.00000464342,0.000015976068,0.0000119728375],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00011723159,0.00049385556,0.0002963179,0.0010228003,0.00082799554,0.00045910935,0.00038871405,0.0007055966,0.011325921],"category_scores_gemma":[0.00024136,0.00018149352,0.000196074,0.0005774111,0.000422184,0.0005933795,0.0006458745,0.00060741126,0.0029055881],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004314651,0.00015821027,0.01151623,0.001201213,0.000029502975,0.0041486984,0.0015334289,0.00043110191,0.02334686,0.0073974663,0.06516307,0.8846427],"study_design_scores_gemma":[0.00007154191,0.00026625645,0.27445817,0.0012139308,0.00004956636,0.012742243,0.00096980494,0.0004720483,0.0020523507,0.0012396905,0.7064191,0.000045441382],"about_ca_topic_score_codex":0.016022569,"about_ca_topic_score_gemma":0.026326455,"teacher_disagreement_score":0.016022569,"about_ca_system_score_codex":0.00077521073,"about_ca_system_score_gemma":0.00037009947,"threshold_uncertainty_score":0.037889004},"labels":[],"label_agreement":null},{"id":"W6949045434","doi":"10.5281/zenodo.12181323","title":"les exercices du petit prof 4e année pdf","year":2024,"lang":"fr","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Industrial management; Executive branch; Context (archaeology)","score_opus":0.02916931783431772,"score_gpt":0.2601560099252759,"score_spread":0.23098669209095815,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949045434","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0016152732,0.0007228383,0.0007477831,0.0030450441,0.003609144,0.00027267876,0.0035498787,0.0012259352,0.9852114],"genre_scores_gemma":[0.0046464605,0.0005967271,0.00049298303,0.0007771501,0.0006258552,0.00012369824,0.0017333494,0.00048770473,0.990516],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990138,0.000097748656,0.000037706028,0.00008553315,0.00057661615,0.00018850497],"domain_scores_gemma":[0.99616265,0.00040885873,0.00013012234,0.00022421265,0.0019520409,0.0011220775],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00093392003,0.00068008754,0.00051948533,0.0015915032,0.0022732227,0.00512769,0.0007831544,0.0016706533,0.76482916],"category_scores_gemma":[0.005136333,0.00040752185,0.0005658859,0.001367607,0.0005490975,0.0031882143,0.003321312,0.0019200669,0.4785252],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000025922809,0.000043129945,0.00023642086,0.000077308156,0.0000011855923,0.00005956697,0.000107534,0.00004074685,0.00023410498,0.0014403594,0.9622607,0.035473034],"study_design_scores_gemma":[0.0000044090507,0.000017040355,0.0010228796,0.00006138464,8.317832e-7,0.000059422935,0.00011862908,0.000018152497,0.000102464786,0.00021225285,0.9983772,0.0000053081053],"about_ca_topic_score_codex":0.003851209,"about_ca_topic_score_gemma":0.011042023,"teacher_disagreement_score":0.23517084,"about_ca_system_score_codex":0.00093870837,"about_ca_system_score_gemma":0.0019256043,"threshold_uncertainty_score":0.33544266},"labels":[],"label_agreement":null},{"id":"W6949096046","doi":"10.5281/zenodo.13335889","title":"·÷±‡±+91-8094774404±‡±÷· Control Your Wife By Vahikaran Mantra","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Wife; Mantra; Miracle; MAGIC (telescope); nobody; Faith","score_opus":0.02165923720022387,"score_gpt":0.25879072300159517,"score_spread":0.2371314858013713,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949096046","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00071374804,0.00028324232,0.00048255938,0.0010600146,0.0013326727,0.00008954464,0.0004149901,0.0016239146,0.99399936],"genre_scores_gemma":[0.0014821775,0.00018247192,0.00024773594,0.00034854253,0.00008511035,0.000026169435,0.00016442555,0.0002517121,0.99721175],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99963796,0.00003568284,0.0000144858595,0.00007426296,0.00015334571,0.000084172796],"domain_scores_gemma":[0.998906,0.00008900208,0.000048535923,0.00006753831,0.00041939196,0.00046949516],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00032945533,0.0007814867,0.0005732704,0.0005587363,0.0017245798,0.0032452445,0.0007656212,0.0012847296,0.942188],"category_scores_gemma":[0.0014820067,0.00033204982,0.0003917556,0.0004655114,0.00034981957,0.0019206683,0.0020221667,0.0016749924,0.90073574],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055341505,0.000066875604,0.00013571278,0.000094303155,0.0000016504708,0.000057123616,0.000072636205,0.000026246396,0.00081610633,0.0008848686,0.89445186,0.10333724],"study_design_scores_gemma":[0.0000094585785,0.000036657617,0.0002846749,0.00005352378,0.000001610327,0.00008589885,0.000115359035,0.000039275747,0.00023703209,0.00013751295,0.99899334,0.0000056951944],"about_ca_topic_score_codex":0.0017196661,"about_ca_topic_score_gemma":0.00344093,"teacher_disagreement_score":0.057811975,"about_ca_system_score_codex":0.00053340703,"about_ca_system_score_gemma":0.00079560326,"threshold_uncertainty_score":0.082461715},"labels":[],"label_agreement":null},{"id":"W6949146885","doi":"10.5281/zenodo.12779153","title":"Transforming Personal Care Packaging: Customized Innovations Redefining Market Dynamics","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Market analysis; Market research; Order (exchange); Market share; Competitive advantage; Identification (biology); Investment (military); Market segmentation; Product (mathematics); Personal care","score_opus":0.021326702296389735,"score_gpt":0.2582462006833947,"score_spread":0.236919498387005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949146885","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.01409342,0.020879926,0.031395257,0.09212581,0.004867479,0.00019098925,0.00029662548,0.0015739297,0.8345765],"genre_scores_gemma":[0.26568142,0.066656806,0.07103922,0.03888739,0.0056905034,0.00032866522,0.00091750367,0.002607549,0.5481909],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968088,0.00079541095,0.00011916841,0.0003675739,0.0015603108,0.0003486568],"domain_scores_gemma":[0.99637115,0.00086318154,0.00021605495,0.0006495767,0.001088921,0.0008111126],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004570468,0.0008510866,0.00028687785,0.0016224062,0.0017770944,0.013405527,0.0013085558,0.004261651,0.052017905],"category_scores_gemma":[0.0070693227,0.00037163318,0.0007853961,0.0020889114,0.0050295494,0.01867277,0.006113601,0.004123927,0.021883033],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003369654,0.000103238825,0.00062575936,0.00030826768,0.0000070217793,0.00025721788,0.0029259496,0.00038160212,0.0024309799,0.57191765,0.16236494,0.2586436],"study_design_scores_gemma":[0.000005109671,0.000035607438,0.00035668767,0.00027012647,0.0000030966994,0.00020517671,0.0011548469,0.0002699521,0.0006122782,0.037737425,0.9593367,0.000012950536],"about_ca_topic_score_codex":0.0021287217,"about_ca_topic_score_gemma":0.0027592161,"teacher_disagreement_score":0.052017905,"about_ca_system_score_codex":0.0038247234,"about_ca_system_score_gemma":0.004557325,"threshold_uncertainty_score":0.17401719},"labels":[],"label_agreement":null},{"id":"W6949202501","doi":"10.5281/zenodo.13125945","title":"NexaSlim Australia Capsules AU, NZ, CA, UK, ZA, FR Weight Loss Supplement For Slim Body Shape","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Weight loss; Overweight; Calorie; Energy expenditure; Stomach tube; Order (exchange); Consumption (sociology); Obesity","score_opus":0.03801023108101807,"score_gpt":0.29342095053154243,"score_spread":0.25541071945052435,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949202501","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005087148,0.009356037,0.0011605645,0.0025100354,0.0024069652,0.000455984,0.0020049391,0.002021088,0.97499734],"genre_scores_gemma":[0.006813818,0.0026547804,0.0009922367,0.000737369,0.000152238,0.000091178095,0.0006215005,0.00025950748,0.9876773],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9998074,0.000022905455,0.0000089472815,0.000036057547,0.00010612527,0.000018664605],"domain_scores_gemma":[0.9996381,0.000051579875,0.000029578701,0.000027075645,0.00011221002,0.00014149035],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00033567369,0.000481554,0.0004043818,0.0007698001,0.0005414777,0.00085395086,0.00046505767,0.0010425722,0.67885303],"category_scores_gemma":[0.0009333294,0.00026283332,0.00036392282,0.00026809226,0.00031018583,0.0010365088,0.001102275,0.0009982545,0.3158281],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00103634,0.00024073462,0.0002950821,0.0009961135,0.0000147746905,0.00021218989,0.00012563277,0.00004951708,0.0098275,0.0026697335,0.5264978,0.45803463],"study_design_scores_gemma":[0.00008249634,0.00019235087,0.0013370682,0.00020953696,0.000007111463,0.00021384686,0.000044012442,0.00006143097,0.0007219662,0.0002451745,0.99687994,0.0000051840175],"about_ca_topic_score_codex":0.0032835447,"about_ca_topic_score_gemma":0.008608147,"teacher_disagreement_score":0.32114697,"about_ca_system_score_codex":0.0004235394,"about_ca_system_score_gemma":0.0006200537,"threshold_uncertainty_score":0.45807713},"labels":[],"label_agreement":null},{"id":"W6949340188","doi":"10.5281/zenodo.14003869","title":"Fausto Romitelli – Professor Bad Trip: Lesson III","year":2024,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Theme (computing); Work (physics); Exposition (narrative); Relation (database); Moment (physics)","score_opus":0.0355656342046667,"score_gpt":0.28653559520664235,"score_spread":0.25096996100197566,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949340188","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0008940747,0.008370599,0.0020751709,0.80783254,0.14083256,0.000052086747,0.0002402601,0.00043179057,0.039270908],"genre_scores_gemma":[0.021553466,0.009257022,0.0028074798,0.4525945,0.056297276,0.00020644326,0.00045014106,0.00085738755,0.45597637],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99874383,0.00034888348,0.000040351424,0.000248868,0.00041846535,0.00019953615],"domain_scores_gemma":[0.9958397,0.00043451643,0.000155613,0.00017317817,0.0010291656,0.0023678178],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028211717,0.00065387134,0.00047731007,0.0005987721,0.0024393909,0.0032273813,0.0008356236,0.0036986961,0.042971775],"category_scores_gemma":[0.012173226,0.00029125647,0.0004531002,0.00044977464,0.0014326144,0.0036171356,0.0036089744,0.0075688125,0.04552753],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000140642705,0.000009106783,0.00008249163,0.000024087154,0.0000014155955,0.000038565457,0.00014671711,0.000015840958,0.000069863534,0.0013787746,0.9904266,0.0077924584],"study_design_scores_gemma":[0.000009433802,0.00002074731,0.00024855623,0.00012737104,0.0000030391916,0.00036736205,0.0004322552,0.000029946039,0.000100925536,0.0021669175,0.99648476,0.000008614923],"about_ca_topic_score_codex":0.0036136894,"about_ca_topic_score_gemma":0.0062501286,"teacher_disagreement_score":0.042971775,"about_ca_system_score_codex":0.0019713372,"about_ca_system_score_gemma":0.0035693352,"threshold_uncertainty_score":0.14375484},"labels":[],"label_agreement":null},{"id":"W6949342783","doi":"10.5281/zenodo.15357562","title":"Dolichogenidea helenedumasae Fernandez-Triana & Boudreault 2025, sp. nov.","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; Parks Canada","funders":"","keywords":"DNA barcoding; Shot (pellet); Ovipositor; Daughter; Cannibalism","score_opus":0.02704380766133067,"score_gpt":0.27819190156279555,"score_spread":0.25114809390146486,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949342783","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.29810014,0.06639983,0.024401095,0.0027088982,0.0032781065,0.0032422363,0.024035309,0.0027830254,0.57505137],"genre_scores_gemma":[0.8367105,0.02471557,0.029547056,0.0028979406,0.0005132486,0.0010272489,0.017826755,0.000252608,0.08650901],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99983525,0.000015386726,0.000014692425,0.00005215051,0.00004421581,0.000038283375],"domain_scores_gemma":[0.9998572,0.000017614893,0.000039905222,0.000017780905,0.000044438762,0.000023061697],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00025883995,0.0013559168,0.0006194486,0.0019965712,0.0014060924,0.0006136078,0.00094398606,0.0008627052,0.013052541],"category_scores_gemma":[0.0006165962,0.00024534797,0.00031052285,0.0015516373,0.0005184974,0.0010673612,0.0007777179,0.001046,0.006834472],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00034367543,0.0001932987,0.0325945,0.00090006465,0.00013881842,0.0027628597,0.00131732,0.0008286018,0.016746484,0.0026329057,0.03878885,0.9027526],"study_design_scores_gemma":[0.00020127036,0.00023402009,0.35084695,0.0016377805,0.00027827543,0.0153317675,0.0022040717,0.0010305259,0.0025998328,0.0014012898,0.6241492,0.00008496851],"about_ca_topic_score_codex":0.046543144,"about_ca_topic_score_gemma":0.07859635,"teacher_disagreement_score":0.046543144,"about_ca_system_score_codex":0.0012140706,"about_ca_system_score_gemma":0.0007726696,"threshold_uncertainty_score":0.092544496},"labels":[],"label_agreement":null},{"id":"W6949619380","doi":"10.5281/zenodo.15241443","title":"Opening up translational data impact through the Data Citation Corpus","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"OpenAlex","funders":"National Center for Advancing Translational Sciences","keywords":"Metadata; Citation; Identifier; Identification (biology); Genomics; Translational research; Profiling (computer programming); Unique identifier; Data curation","score_opus":0.12382795325374887,"score_gpt":0.3570606096248361,"score_spread":0.23323265637108723,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949619380","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.019559054,0.0033236123,0.010815636,0.003299887,0.0007587041,0.0005940536,0.93971884,0.0026705747,0.019259566],"genre_scores_gemma":[0.06315193,0.003052796,0.048106864,0.0007997937,0.00069036864,0.0046980362,0.87318784,0.0022488257,0.0040635667],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.96837896,0.006687012,0.008588947,0.004594446,0.010513869,0.0012366781],"domain_scores_gemma":[0.76413274,0.16495973,0.022928242,0.019805156,0.025490794,0.0026832535],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.022186473,0.0010841504,0.0019006078,0.10721412,0.002859328,0.012091394,0.0023697973,0.0018806421,0.021008613],"category_scores_gemma":[0.18862282,0.00089351076,0.0011486833,0.14053723,0.002210974,0.007883946,0.010164585,0.0025982056,0.0076454543],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040547363,0.00015074981,0.0423201,0.012551713,0.0004029356,0.00063697796,0.0077500185,0.0031039312,0.00301466,0.08634988,0.67224646,0.17106713],"study_design_scores_gemma":[0.00009267996,0.00004333676,0.037011508,0.0029309255,0.0001528247,0.00019759151,0.0026858863,0.0022224342,0.0017563431,0.024302857,0.9284346,0.00016898125],"about_ca_topic_score_codex":0.011160957,"about_ca_topic_score_gemma":0.014741787,"teacher_disagreement_score":0.9976302,"about_ca_system_score_codex":0.0046127304,"about_ca_system_score_gemma":0.008012157,"threshold_uncertainty_score":0.117334664},"labels":[],"label_agreement":null},{"id":"W6949688685","doi":"10.5281/zenodo.15357543","title":"Dolichogenidea felipechavarriai Fernandez-Triana & Boudreault 2025, sp. nov.","year":2025,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Guelph; Parks Canada","funders":"","keywords":"DNA barcoding; Voucher; Barcode; Type (biology); Margin (machine learning); Race (biology)","score_opus":0.02678463678136255,"score_gpt":0.2784801310007726,"score_spread":0.25169549421941007,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949688685","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15743755,0.046366706,0.01848362,0.0033550153,0.0033436986,0.0031891128,0.028449986,0.0026979423,0.73667634],"genre_scores_gemma":[0.78848994,0.033136073,0.038620263,0.0047161654,0.0008878432,0.0015854893,0.021607269,0.00042486071,0.11053211],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.999736,0.000015839267,0.000023695156,0.00007502815,0.00009481318,0.00005448454],"domain_scores_gemma":[0.9997826,0.000018245613,0.000061437815,0.000029675406,0.0000776435,0.000030449562],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002568428,0.0014139782,0.0006704914,0.0023026054,0.0019300127,0.0008475532,0.001109623,0.0010058196,0.013350192],"category_scores_gemma":[0.0008011494,0.00037542733,0.00038116073,0.0018853327,0.0007541199,0.0013139222,0.0009850191,0.0015764972,0.007372585],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00028501655,0.00021456803,0.04914316,0.0009129932,0.00019709616,0.0025898207,0.0012336652,0.001311879,0.006699467,0.0039996793,0.08756021,0.84585255],"study_design_scores_gemma":[0.0001391414,0.0001116527,0.34175262,0.0019712506,0.00033626927,0.0131562045,0.002033374,0.0015230255,0.0021703732,0.001989917,0.6347309,0.000085369444],"about_ca_topic_score_codex":0.094201095,"about_ca_topic_score_gemma":0.14124948,"teacher_disagreement_score":0.094201095,"about_ca_system_score_codex":0.002152062,"about_ca_system_score_gemma":0.0014048765,"threshold_uncertainty_score":0.18730557},"labels":[],"label_agreement":null},{"id":"W6949708540","doi":"10.5281/zenodo.4474154","title":"Zagrammosoma americanum Girault","year":2021,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Holotype; Wing; Perpendicular; Occiput; Dorsum; Vertex (graph theory)","score_opus":0.02716063941944543,"score_gpt":0.2570523258108584,"score_spread":0.22989168639141297,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6949708540","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8274706,0.0065340614,0.0020284303,0.0004107407,0.00020362496,0.00013534394,0.0011901781,0.00027613115,0.16175096],"genre_scores_gemma":[0.9874564,0.0010841134,0.001287987,0.00016626378,0.00003865673,0.000023046123,0.00057021284,0.000009339809,0.009363994],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99993885,0.000007239132,0.000004659973,0.000021612634,0.000017474056,0.000010124856],"domain_scores_gemma":[0.9999615,0.000004774576,0.00001488216,0.0000049525306,0.00000825216,0.0000056529743],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000055347457,0.00027223714,0.00016306492,0.000858352,0.00057328265,0.00022008574,0.00023043399,0.00020296218,0.0052216463],"category_scores_gemma":[0.00011068028,0.000100775964,0.00010076107,0.00047632537,0.00029798743,0.00035949252,0.00047033455,0.00030262352,0.001110925],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007353401,0.00024356646,0.10768028,0.0011449121,0.00011408428,0.012347604,0.0033730841,0.002306537,0.19243607,0.016447129,0.020804824,0.6423665],"study_design_scores_gemma":[0.000095036165,0.00020544202,0.6193483,0.00027931028,0.0001110617,0.01838214,0.0017696593,0.00094004464,0.006322541,0.0032863922,0.349222,0.00003804647],"about_ca_topic_score_codex":0.0052652857,"about_ca_topic_score_gemma":0.012426798,"teacher_disagreement_score":0.0052652857,"about_ca_system_score_codex":0.00028145278,"about_ca_system_score_gemma":0.0001221124,"threshold_uncertainty_score":0.017468095},"labels":[],"label_agreement":null},{"id":"W6958006325","doi":"10.6084/m9.figshare.12799362.v1","title":"Additional file 2 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2020,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Disease; Rare events","score_opus":0.06574014444580477,"score_gpt":0.3213320825948276,"score_spread":0.2555919381490228,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958006325","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00009269116,0.00010632021,0.0005217211,0.00020441347,0.000027190768,0.00039367998,0.99742115,0.00036339267,0.00086938665],"genre_scores_gemma":[0.004527083,0.0010080058,0.023168553,0.0018574223,0.00020479466,0.017072668,0.94003594,0.0017008396,0.01042476],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99703133,0.00074621325,0.0011198303,0.00053356116,0.00037869587,0.0001903852],"domain_scores_gemma":[0.90955496,0.07578852,0.005682785,0.0025756024,0.0052308603,0.0011672345],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006247852,0.0012252172,0.0019922128,0.009464934,0.00085431465,0.0023921237,0.0018859636,0.0016085759,0.7550236],"category_scores_gemma":[0.06694781,0.0008717738,0.0023873604,0.010701838,0.00047061619,0.0025420056,0.0023384632,0.0013030104,0.08799909],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042570653,0.000049676742,0.0013349056,0.06305469,0.00030075535,0.00008805862,0.00016110652,0.0003824191,0.00031184423,0.0020785471,0.91538906,0.016423237],"study_design_scores_gemma":[0.00349714,0.00014138801,0.008029251,0.02026518,0.0009505265,0.00029937323,0.00029407683,0.0006173808,0.0007943597,0.012278306,0.9526897,0.0001433253],"about_ca_topic_score_codex":0.004655293,"about_ca_topic_score_gemma":0.014883402,"teacher_disagreement_score":0.9937521,"about_ca_system_score_codex":0.0015058094,"about_ca_system_score_gemma":0.0060192463,"threshold_uncertainty_score":0.34942907},"labels":[],"label_agreement":null},{"id":"W6958049890","doi":"10.6084/m9.figshare.20175983.v1","title":"Additional file 5 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07351708511160826,"score_gpt":0.3383399881375442,"score_spread":0.2648229030259359,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958049890","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0001347019,0.000019795243,0.0011646077,0.00014578186,0.000056248366,0.00005909576,0.99410725,0.0017964502,0.002516078],"genre_scores_gemma":[0.00455788,0.00018329252,0.012237526,0.0004619074,0.00008408871,0.00075177377,0.9679042,0.0043514706,0.009467851],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9994585,0.00008380753,0.000093830444,0.00013664884,0.00013853837,0.000088603185],"domain_scores_gemma":[0.98918766,0.007705075,0.0004886661,0.00073060533,0.0015158842,0.0003721882],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001667399,0.0011218827,0.0010241884,0.0036949692,0.00084620074,0.002465873,0.0018418308,0.0012568136,0.8065459],"category_scores_gemma":[0.015950842,0.0007313476,0.0013691474,0.005136528,0.0004570945,0.0034133254,0.0017602218,0.001425894,0.2418952],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018394343,0.000053676282,0.00069694314,0.0016221808,0.000034062345,0.00006212808,0.00009509636,0.0003387428,0.00025250964,0.0023640743,0.98576975,0.008526906],"study_design_scores_gemma":[0.0009962054,0.000038643917,0.0037596368,0.00086351694,0.000078361445,0.00023455257,0.0002478839,0.0012191631,0.0012622568,0.013672169,0.9775464,0.000081172846],"about_ca_topic_score_codex":0.008616981,"about_ca_topic_score_gemma":0.01130993,"teacher_disagreement_score":0.8065459,"about_ca_system_score_codex":0.0015989777,"about_ca_system_score_gemma":0.0023749673,"threshold_uncertainty_score":0.2759388},"labels":[],"label_agreement":null},{"id":"W6958120276","doi":"10.6084/m9.figshare.12799362","title":"Additional file 2 of A guide to writing systematic reviews of rare disease treatments to generate FAIR-compliant datasets: building a Treatabolome","year":2021,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa; Ottawa Hospital; Children's Hospital of Eastern Ontario","funders":"","keywords":"Systematic review; Rare disease; MEDLINE; Disease; Rare events","score_opus":0.0504202911666707,"score_gpt":0.3255831022755051,"score_spread":0.2751628111088344,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958120276","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00009401381,0.00010598268,0.0005278739,0.00020400817,0.000027278396,0.00039255334,0.99741817,0.00036681045,0.0008632394],"genre_scores_gemma":[0.00454858,0.0009984785,0.023377394,0.0018352075,0.0002027333,0.016971212,0.9400198,0.0017095887,0.010336977],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9970356,0.0007456552,0.001118262,0.0005345648,0.0003762661,0.00018952263],"domain_scores_gemma":[0.9097334,0.07565419,0.0056690024,0.002558743,0.0052251425,0.0011594697],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.006236715,0.0012268733,0.0019873013,0.009423827,0.0008518359,0.0023936399,0.0018741533,0.0015980371,0.75149596],"category_scores_gemma":[0.066778585,0.0008688528,0.0023878135,0.010654296,0.0004713089,0.0025234397,0.0023248561,0.0012994077,0.086941466],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042716903,0.000049928927,0.0013583608,0.06318174,0.00030385822,0.00008912907,0.00016266716,0.0003877754,0.000316053,0.0020857041,0.9151549,0.016482767],"study_design_scores_gemma":[0.0034698474,0.00014093402,0.008048594,0.020075906,0.00095346716,0.00029814168,0.00029596125,0.00062398403,0.00080140436,0.012275011,0.95287377,0.00014299825],"about_ca_topic_score_codex":0.004631547,"about_ca_topic_score_gemma":0.0147834625,"teacher_disagreement_score":0.75149596,"about_ca_system_score_codex":0.001504551,"about_ca_system_score_gemma":0.0059745926,"threshold_uncertainty_score":0.35446084},"labels":[],"label_agreement":null},{"id":"W6958159083","doi":"10.60692/mce4p-kqj86","title":"A Comprehensive Health Effects Assessment of the Use of Sanitizers and Disinfectants during COVID-19 Pandemic: A Global Survey","year":2022,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Alberta; Health Sciences North","funders":"","keywords":"Human health; Disinfectant; Eye irritation; Personal protective equipment; Global health; Hand sanitizer; Risk assessment","score_opus":0.10100746484226124,"score_gpt":0.31324001658850636,"score_spread":0.2122325517462451,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958159083","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9956477,0.0004552164,0.00017831437,0.00006072073,0.0000052521277,0.0001007847,0.0028623708,0.000006877322,0.00068276044],"genre_scores_gemma":[0.9974074,0.00037715037,0.00033254994,0.00006308056,0.000007016776,0.00007770482,0.0014993481,0.0000016622462,0.00023407719],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99938464,0.00019229572,0.00012328503,0.00009865522,0.0001242995,0.00007683943],"domain_scores_gemma":[0.998228,0.00023542545,0.00092916086,0.00011708982,0.0002779503,0.00021234866],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014000962,0.00029181206,0.00030746768,0.0013223366,0.00027665042,0.0003711973,0.0002647013,0.0005973243,0.0014255194],"category_scores_gemma":[0.0015511845,0.00037237792,0.0008279825,0.001365963,0.00026186506,0.00060354196,0.0007317596,0.00048066152,0.00027279602],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000033574863,0.000064803586,0.9973792,0.00007865665,0.00006385283,0.000032484662,0.0001612113,0.000055370696,0.00019466355,0.000008095726,0.00011093974,0.0018171835],"study_design_scores_gemma":[0.00000249626,0.00017913566,0.9991725,0.00001262821,0.000017932018,0.00006546954,0.0002579844,0.00006124482,0.000037156748,0.0000047975,0.00018542884,0.000003245804],"about_ca_topic_score_codex":0.0035357433,"about_ca_topic_score_gemma":0.0035133327,"teacher_disagreement_score":0.0035357433,"about_ca_system_score_codex":0.00029945927,"about_ca_system_score_gemma":0.00028818083,"threshold_uncertainty_score":0.007404506},"labels":[],"label_agreement":null},{"id":"W6958487928","doi":"10.6084/m9.figshare.25302906","title":"Additional file 1 of GPAD: a natural language processing-based application to extract the gene-disease association discovery information from OMIM","year":2024,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Dependency (UML); Association (psychology); Natural language; Information extraction; Feature (linguistics)","score_opus":0.006152937457157603,"score_gpt":0.2413419467694597,"score_spread":0.23518900931230208,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958487928","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00018592885,0.000020169002,0.0012227154,0.00008334703,0.000027305034,0.00007096212,0.99568594,0.0017320865,0.0009715208],"genre_scores_gemma":[0.0053947717,0.00012776307,0.010509148,0.0004054339,0.00008970071,0.001333194,0.9741344,0.0028125602,0.005193094],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9993261,0.0001218445,0.000116916184,0.00019182914,0.00016483266,0.00007851905],"domain_scores_gemma":[0.9868973,0.010365432,0.0005650438,0.00067939184,0.0011459605,0.00034688687],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0018928698,0.0015197471,0.0011616511,0.0030709154,0.0007057411,0.0017209493,0.0021320619,0.0011567147,0.77387583],"category_scores_gemma":[0.018499603,0.0007948213,0.0011306233,0.0037024796,0.00040315356,0.0017540667,0.0016924066,0.0012207888,0.19192666],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039080656,0.00009485665,0.0018778188,0.0032269196,0.000068182155,0.00013827698,0.00008190899,0.0007012983,0.0005706881,0.0012289261,0.9773791,0.014241293],"study_design_scores_gemma":[0.0028924514,0.00026194708,0.012527595,0.0015145226,0.00024490405,0.0007500248,0.00027858472,0.0039981664,0.0045679463,0.014474932,0.9583049,0.00018403651],"about_ca_topic_score_codex":0.0039660903,"about_ca_topic_score_gemma":0.0062069665,"teacher_disagreement_score":0.77387583,"about_ca_system_score_codex":0.0010681562,"about_ca_system_score_gemma":0.0018117997,"threshold_uncertainty_score":0.32253867},"labels":[],"label_agreement":null},{"id":"W6958768399","doi":"10.6084/m9.figshare.c.6836253","title":"Incidence and prevalence of neurofibromatosis type 1 and 2: a systematic review and meta-analysis","year":2023,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University Health Network; University of Toronto","funders":"","keywords":"Incidence (geometry); Prevalence; Neurofibromatosis; Epidemiology; Meta-analysis; Medical record; Systematic review","score_opus":0.05512817342363259,"score_gpt":0.30862370440759057,"score_spread":0.253495530983958,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6958768399","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002213041,0.99597365,0.00044719412,0.00020067324,0.00014826268,0.00026804843,0.00056239316,0.000018464394,0.00016821235],"genre_scores_gemma":[0.0922312,0.9003361,0.0032200248,0.00086945185,0.00033222281,0.0017148507,0.0010233739,0.00002902136,0.00024378624],"study_design_codex":"systematic_review","study_design_gemma":"meta_analysis","domain_scores_codex":[0.97953904,0.009129226,0.0064239544,0.0018969208,0.00252427,0.0004866008],"domain_scores_gemma":[0.96337867,0.024703741,0.0069507244,0.0010261465,0.003619325,0.0003214543],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.026823705,0.0034934047,0.021634532,0.011474102,0.00058588997,0.0045546065,0.0023377736,0.002433039,0.004107151],"category_scores_gemma":[0.06691832,0.0020625552,0.03344407,0.010932903,0.0007552876,0.0028656486,0.0018326804,0.0021466315,0.00034479712],"study_design_candidate":"meta_analysis","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00076705863,0.000019792911,0.0040598847,0.52796614,0.4536954,0.00012292142,0.00008406931,0.00044079928,0.00013395588,0.00014067552,0.00089362747,0.011675687],"study_design_scores_gemma":[0.0004728407,0.00016101071,0.0056058047,0.077675715,0.9130492,0.00021842227,0.00006417282,0.0001948936,0.00010669576,0.00025420325,0.0021626768,0.000034405653],"about_ca_topic_score_codex":0.0070080115,"about_ca_topic_score_gemma":0.0156035125,"teacher_disagreement_score":0.026823705,"about_ca_system_score_codex":0.0036922917,"about_ca_system_score_gemma":0.0047498024,"threshold_uncertainty_score":0.141859},"labels":[],"label_agreement":null},{"id":"W6961810291","doi":"10.15468/dl.etgzfy","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6961810291","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008538772,0.00004756773,0.00007823574,0.00005344925,0.000014995826,0.000008491959,0.9980563,0.0007808304,0.00087473466],"genre_scores_gemma":[0.00018432249,0.00004001172,0.00028192936,0.000046416535,0.0000028495547,0.00003687378,0.9988286,0.00014582263,0.00043313042],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988926,0.00015016847,0.00015189209,0.00038390476,0.00026306367,0.0001583263],"domain_scores_gemma":[0.9978695,0.00058429246,0.00020000798,0.00055639935,0.00052165723,0.000268185],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00096627726,0.0024305787,0.0015123528,0.004727095,0.0010375016,0.0024250988,0.0030894834,0.00205377,0.10177354],"category_scores_gemma":[0.0056903297,0.0008820091,0.0012939544,0.0086971205,0.0004734676,0.0022409053,0.00249195,0.0019523342,0.15490742],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004012236,0.000017346452,0.00043307693,0.00054316723,0.000016747657,0.000021760583,0.00002535382,0.00018167408,0.00014769379,0.0004781422,0.9962186,0.0018762893],"study_design_scores_gemma":[0.000085174484,0.000011486259,0.0017166425,0.00016925808,0.000015428262,0.000060443275,0.00007413346,0.00030999578,0.000276972,0.0009543936,0.99630666,0.000019472041],"about_ca_topic_score_codex":0.02430223,"about_ca_topic_score_gemma":0.042368833,"teacher_disagreement_score":0.89822644,"about_ca_system_score_codex":0.0017888953,"about_ca_system_score_gemma":0.002638144,"threshold_uncertainty_score":0.34046638},"labels":[],"label_agreement":null},{"id":"W6961868853","doi":"10.15468/dl.bdcv5g","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Matching (statistics); Range (aeronautics); Set (abstract data type); Identification (biology); Download","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6961868853","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008959507,0.000050924187,0.00007618076,0.00005521788,0.000015171976,0.000009144223,0.9979882,0.0007956916,0.00091994967],"genre_scores_gemma":[0.0001935248,0.00004129645,0.0002935007,0.000047555684,0.0000029584069,0.00004045767,0.99878794,0.0001472832,0.00044551215],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99889404,0.00014742462,0.00015034307,0.0003807111,0.00026867318,0.00015882438],"domain_scores_gemma":[0.9978498,0.0006083812,0.00020104667,0.0005460857,0.0005242273,0.00027051],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00096908875,0.0025187188,0.0015090484,0.0049632764,0.0010823709,0.002458874,0.00310992,0.0021830664,0.0990179],"category_scores_gemma":[0.0057622236,0.00089841516,0.0013074876,0.008964359,0.00048517314,0.0022424958,0.002532537,0.001982235,0.1503856],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000040546354,0.000017951303,0.00044098657,0.00057793735,0.000017071296,0.000024190747,0.00002675785,0.00018316774,0.00015452849,0.00046974438,0.99612147,0.0019256321],"study_design_scores_gemma":[0.00008766425,0.000011617341,0.0017708725,0.00017789136,0.000016057638,0.00006348999,0.00007554788,0.00031425344,0.0002861889,0.0009525285,0.9962238,0.000020147649],"about_ca_topic_score_codex":0.024676505,"about_ca_topic_score_gemma":0.042651404,"teacher_disagreement_score":0.9009821,"about_ca_system_score_codex":0.0018219461,"about_ca_system_score_gemma":0.0026737815,"threshold_uncertainty_score":0.33124787},"labels":[],"label_agreement":null},{"id":"W6962261565","doi":"10.15468/dl.njmg8m","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Data set","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6962261565","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007466049,0.000045203367,0.000057482048,0.00004849035,0.000011892254,0.000008123893,0.9984175,0.0006085897,0.00072800624],"genre_scores_gemma":[0.00018277587,0.000037507838,0.00020261353,0.000044994165,0.0000026970563,0.000033619886,0.9990195,0.00012014659,0.00035613793],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988795,0.0001525784,0.0001538232,0.00037723762,0.00026807497,0.00016882915],"domain_scores_gemma":[0.9978878,0.00056052656,0.0002109814,0.0005608155,0.00049165596,0.0002881338],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010653888,0.0023998441,0.0015727669,0.005596152,0.0009888663,0.0025541114,0.0029739097,0.0022145957,0.10047041],"category_scores_gemma":[0.0056003914,0.0008795865,0.0012345944,0.009126935,0.00048249547,0.0022266982,0.0026109682,0.0020041915,0.16239941],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000046739755,0.000017893282,0.00044326612,0.0006858186,0.000020391493,0.000026791691,0.000026148686,0.00018181815,0.00019124777,0.00047322226,0.9959318,0.0019548878],"study_design_scores_gemma":[0.00009339124,0.000011026757,0.0019611297,0.00020382558,0.000017309549,0.000063159445,0.00006932462,0.0002807299,0.0003134846,0.0009153167,0.99605227,0.00001905202],"about_ca_topic_score_codex":0.019748729,"about_ca_topic_score_gemma":0.032718383,"teacher_disagreement_score":0.8995296,"about_ca_system_score_codex":0.001807123,"about_ca_system_score_gemma":0.0024714284,"threshold_uncertainty_score":0.33610702},"labels":[],"label_agreement":null},{"id":"W6962276807","doi":"10.15468/dl.mmw7w4","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Set (abstract data type); Data set","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6962276807","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00007719971,0.000043538912,0.000068112844,0.000057707803,0.000013684196,0.000008772004,0.9981225,0.0007542637,0.00085427717],"genre_scores_gemma":[0.00016923902,0.000035471418,0.00023102887,0.000045012952,0.0000028123652,0.000037166243,0.998961,0.00013271884,0.0003856808],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989209,0.00013854416,0.00013797535,0.00036894533,0.0002684284,0.00016524293],"domain_scores_gemma":[0.99787045,0.0005770102,0.00018100278,0.0005682178,0.00053258357,0.00027074217],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.000981806,0.0024982553,0.001600799,0.0049209995,0.0011620589,0.0025391744,0.0033409311,0.0023463503,0.10674568],"category_scores_gemma":[0.005567836,0.0009290689,0.0013353658,0.009203633,0.0005046913,0.0024235926,0.002587172,0.0021845242,0.17491797],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003687325,0.000016636886,0.00037854695,0.0004895637,0.0000151952845,0.00002221822,0.000023865765,0.00015560846,0.00013065587,0.00040817424,0.9967314,0.001591199],"study_design_scores_gemma":[0.000097874385,0.000010682771,0.0018050798,0.00016963124,0.000015313604,0.0000606404,0.00008192534,0.00030946237,0.00027596747,0.0009414952,0.99621177,0.000020058833],"about_ca_topic_score_codex":0.027370173,"about_ca_topic_score_gemma":0.046399206,"teacher_disagreement_score":0.89325434,"about_ca_system_score_codex":0.0019290698,"about_ca_system_score_gemma":0.0026955262,"threshold_uncertainty_score":0.3570999},"labels":[],"label_agreement":null},{"id":"W6962749680","doi":"10.15468/dl.esa3fg","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Set (abstract data type)","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6962749680","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008166536,0.000043986016,0.000068384456,0.00006333358,0.000014932841,0.000008928055,0.9982015,0.00070694217,0.0008102347],"genre_scores_gemma":[0.00017915557,0.000035489033,0.00023718056,0.00004692077,0.0000029270716,0.00003732778,0.9989568,0.000116948686,0.00038712996],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.998882,0.00015294213,0.00014771867,0.0003816393,0.00026991425,0.00016568227],"domain_scores_gemma":[0.99774987,0.0006150303,0.00019740095,0.00062031526,0.00054044987,0.00027688718],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009997274,0.0023926378,0.0015471646,0.004709861,0.0011621785,0.0025193708,0.0032637476,0.0023448395,0.09902594],"category_scores_gemma":[0.0057583135,0.000880701,0.0013158012,0.008680995,0.0005096518,0.0023809727,0.0026106974,0.0021500324,0.15870894],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000038590857,0.000017417819,0.0004037293,0.0005024749,0.000015691987,0.000022347438,0.000024823561,0.00016770109,0.00012447413,0.00042798676,0.9966445,0.0016102013],"study_design_scores_gemma":[0.00010427863,0.000011995936,0.0019506308,0.00017880107,0.00001594474,0.00006741035,0.00008922461,0.00032211002,0.00026973552,0.0009836766,0.99598515,0.000021134474],"about_ca_topic_score_codex":0.026789932,"about_ca_topic_score_gemma":0.046157666,"teacher_disagreement_score":0.90097404,"about_ca_system_score_codex":0.0018707357,"about_ca_system_score_gemma":0.0027081268,"threshold_uncertainty_score":0.33127475},"labels":[],"label_agreement":null},{"id":"W6967233829","doi":"10.5281/zenodo.11400056","title":"Puravive Affiliate Program: Ingredients, Pros, Cons, Benefits, Side Effects, Customer Reviews!","year":2024,"lang":"en","type":"other","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Limiting; Work (physics); Filter (signal processing); Natural (archaeology); Population; Quality (philosophy)","score_opus":0.02751853933957446,"score_gpt":0.28369528586603593,"score_spread":0.2561767465264615,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6967233829","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0020038863,0.008020684,0.004782575,0.0051720007,0.0037010368,0.0017754787,0.022016505,0.037486643,0.9150411],"genre_scores_gemma":[0.004585387,0.0048388285,0.0049686613,0.005150356,0.0015277233,0.0011069984,0.008491517,0.006267953,0.9630625],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989599,0.00018356639,0.00007232663,0.0001162851,0.0005851196,0.00008293712],"domain_scores_gemma":[0.9951566,0.0010916716,0.0003659302,0.0004260861,0.0019397421,0.001019845],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.002038891,0.0012814788,0.0014034149,0.0021249065,0.0010255128,0.003972896,0.0016993927,0.0028327547,0.80101556],"category_scores_gemma":[0.008171948,0.0008373106,0.0009370901,0.0015626894,0.0004086134,0.0044107437,0.0021048286,0.0020874597,0.7682016],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000750002,0.000071168084,0.000040301064,0.0004002861,0.00000487745,0.0000268965,0.000020314861,0.000018088014,0.0008167709,0.0002976091,0.92448944,0.07373923],"study_design_scores_gemma":[0.000060052036,0.00008137175,0.0002847992,0.00013224212,0.000009454014,0.000091892885,0.000023741562,0.00007685505,0.00068675564,0.00023305853,0.99830365,0.000016118547],"about_ca_topic_score_codex":0.0011794132,"about_ca_topic_score_gemma":0.0029305958,"teacher_disagreement_score":0.19898444,"about_ca_system_score_codex":0.0006767397,"about_ca_system_score_gemma":0.001572019,"threshold_uncertainty_score":0.2838272},"labels":[],"label_agreement":null},{"id":"W6969083413","doi":"10.5281/zenodo.6209898","title":"Procladius (Holotanypus) sublettei Roback","year":2010,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Buoy; Beaver; Fishing; Government (linguistics)","score_opus":0.020935558439685056,"score_gpt":0.2481087032907976,"score_spread":0.22717314485111254,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6969083413","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.85991216,0.004634641,0.0036086533,0.0002002563,0.00023457414,0.00018218682,0.0013303059,0.0003458245,0.12955156],"genre_scores_gemma":[0.9620007,0.0016137164,0.004825144,0.00027755715,0.00007083802,0.00021744675,0.001849267,0.000045417557,0.029099772],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9999037,0.000009954131,0.000009117601,0.000037753183,0.00002415879,0.000015364112],"domain_scores_gemma":[0.9997396,0.00005135147,0.00010280706,0.000022746042,0.00005305339,0.000030457719],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000102293125,0.00040439088,0.00019706077,0.0009894157,0.001046307,0.00032101487,0.00042188325,0.0002514549,0.008738166],"category_scores_gemma":[0.00034050824,0.0002245255,0.00012387699,0.00037388297,0.00044452187,0.00048688153,0.00073572877,0.00030158434,0.0030227064],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.001811506,0.00026479815,0.17827773,0.001945915,0.00008209341,0.0033901534,0.006228852,0.000865203,0.0826552,0.006325401,0.022419319,0.69573385],"study_design_scores_gemma":[0.00009055569,0.00070453755,0.767814,0.0005362733,0.000090071044,0.006043392,0.0024296753,0.0006844331,0.007086011,0.00067544007,0.2137928,0.000052906424],"about_ca_topic_score_codex":0.012604684,"about_ca_topic_score_gemma":0.029038753,"teacher_disagreement_score":0.012604684,"about_ca_system_score_codex":0.00047824142,"about_ca_system_score_gemma":0.00029821566,"threshold_uncertainty_score":0.029232085},"labels":[],"label_agreement":null},{"id":"W6969088978","doi":"10.5281/zenodo.4609335","title":"CINECA_Query expansion service_D1.2","year":2020,"lang":"en","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Hospital for Sick Children","funders":"European Commission","keywords":"Discoverability; Ontology; Representation (politics); SPARQL; Ranking (information retrieval); External Data Representation; Data access; Data integration; RDF","score_opus":0.04041229772472648,"score_gpt":0.2513374146380585,"score_spread":0.210925116913332,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6969088978","genre_codex":"software","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0033730287,0.00049677474,0.17106208,0.003612976,0.0005946871,0.0012195127,0.08486801,0.6825712,0.052201763],"genre_scores_gemma":[0.10518052,0.0014290293,0.25979754,0.014338276,0.0008163243,0.0034854857,0.39189553,0.15372065,0.06933663],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99340516,0.0013463778,0.00069827883,0.0014037655,0.002573255,0.0005732625],"domain_scores_gemma":[0.98564756,0.004785909,0.00037646346,0.005327173,0.0031258503,0.00073706947],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008846716,0.002734361,0.0018215029,0.0032779905,0.0016162059,0.0074059544,0.004547221,0.0030415852,0.07916607],"category_scores_gemma":[0.024321146,0.001711697,0.0031265526,0.0029796087,0.0015626015,0.00930788,0.009814223,0.004266407,0.052178543],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010416937,0.00016762903,0.0030512859,0.00062307995,0.00013626272,0.00042936607,0.000833954,0.0015292775,0.00592775,0.025059037,0.90123975,0.059960935],"study_design_scores_gemma":[0.0003653976,0.000089381996,0.002691624,0.00017622994,0.000044982775,0.00048067406,0.00039395972,0.038871706,0.013253939,0.021196805,0.922245,0.00019028704],"about_ca_topic_score_codex":0.033549044,"about_ca_topic_score_gemma":0.017307887,"teacher_disagreement_score":0.07916607,"about_ca_system_score_codex":0.0035089734,"about_ca_system_score_gemma":0.0045408043,"threshold_uncertainty_score":0.2648369},"labels":[],"label_agreement":null},{"id":"W6976573286","doi":"10.60692/q06w7-hfy72","title":"An ecological study of the association between environmental indicators and early childhood caries","year":2020,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Manitoba","funders":"","keywords":"Ecological study; Emission intensity; Ecosystem; Nitrous oxide; Intensity (physics); Association (psychology); Carbon dioxide; Index (typography)","score_opus":0.013865237742031033,"score_gpt":0.1968739831099636,"score_spread":0.18300874536793257,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976573286","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99925655,0.0001955285,0.00013857988,0.00003229399,0.0000013112298,0.0000021794315,0.00008634518,0.0000013200779,0.0002858445],"genre_scores_gemma":[0.9996896,0.00007447582,0.00013871345,0.0000046996333,0.0000018665953,0.0000015440713,0.000044847235,4.6125487e-7,0.000043806594],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9989035,0.00063183595,0.000073676994,0.00012898889,0.0001473341,0.000114609495],"domain_scores_gemma":[0.99645317,0.0015648557,0.0011537482,0.00014399903,0.00038807525,0.00029609218],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015471802,0.00018586637,0.00022766965,0.0019941924,0.00045185798,0.00058905943,0.0001705235,0.00018041601,0.0009338345],"category_scores_gemma":[0.004277825,0.00018111597,0.0004021905,0.0027802333,0.00044158805,0.00044547272,0.0007290106,0.00038205707,0.00006449447],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003691237,0.000030914824,0.996485,0.000020504564,0.00008705697,0.000055553253,0.0002402613,0.00007674359,0.00016738554,0.00007203221,0.000030025776,0.0026975241],"study_design_scores_gemma":[7.265992e-7,0.000026513808,0.99928576,0.0000039248953,0.000013789186,0.000048995847,0.00036618445,0.000120535464,0.000028001865,0.000034053784,0.00007033011,0.0000012790847],"about_ca_topic_score_codex":0.012963783,"about_ca_topic_score_gemma":0.021596478,"teacher_disagreement_score":0.012963783,"about_ca_system_score_codex":0.00039202685,"about_ca_system_score_gemma":0.0005786396,"threshold_uncertainty_score":0.025776625},"labels":[],"label_agreement":null},{"id":"W6976865229","doi":"10.6084/m9.figshare.20175986.v1","title":"Additional file 6 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07351232016756304,"score_gpt":0.3383119798651275,"score_spread":0.26479965969756447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976865229","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00012579093,0.000019057341,0.0011052993,0.00014212106,0.00005683083,0.000056591725,0.9941446,0.0018084724,0.0025412722],"genre_scores_gemma":[0.0043135225,0.00018426261,0.012015151,0.00044969222,0.000086063854,0.0007275076,0.96800214,0.004564379,0.009657185],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99945134,0.00008377704,0.00009438495,0.00013902257,0.00014218022,0.00008934551],"domain_scores_gemma":[0.9895538,0.007353454,0.00047066007,0.00074381905,0.001509138,0.00036909466],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016709118,0.0011292972,0.0010473955,0.003704612,0.00082991185,0.002490701,0.0018611593,0.0012802711,0.81457025],"category_scores_gemma":[0.01572349,0.000748244,0.0014034029,0.005031253,0.00045605673,0.0034576845,0.001768972,0.0014673845,0.26104245],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017135222,0.000052610514,0.00064526853,0.0015329907,0.000033017273,0.000058922327,0.00009097645,0.00031326155,0.00022958628,0.0021470312,0.9866666,0.008058523],"study_design_scores_gemma":[0.0010140722,0.000037135615,0.0036244614,0.000821176,0.00007511889,0.00022247696,0.00023813701,0.0011731541,0.0011655509,0.013038022,0.9785091,0.00008145008],"about_ca_topic_score_codex":0.008592859,"about_ca_topic_score_gemma":0.011564062,"teacher_disagreement_score":0.81457025,"about_ca_system_score_codex":0.0015505982,"about_ca_system_score_gemma":0.002367243,"threshold_uncertainty_score":0.26449305},"labels":[],"label_agreement":null},{"id":"W6976877773","doi":"10.6084/m9.figshare.19759380.v1","title":"Additional file 2 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test)","score_opus":0.06545233424165164,"score_gpt":0.4051565814282834,"score_spread":0.3397042471866318,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976877773","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002811582,0.0000147255105,0.00095326535,0.00010883958,0.000022619837,0.00008766399,0.9964204,0.001310295,0.0008011109],"genre_scores_gemma":[0.006328249,0.0000789042,0.009770464,0.0003094668,0.00009348639,0.0011588796,0.97611994,0.0012186328,0.004921907],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99911124,0.00013378268,0.00022478885,0.00024366229,0.00020800358,0.00007865346],"domain_scores_gemma":[0.9826717,0.012227457,0.0012287757,0.0009833799,0.002499627,0.00038921068],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015967397,0.0010129962,0.000758618,0.003278639,0.000535015,0.0014578437,0.0014361892,0.0008755269,0.6593007],"category_scores_gemma":[0.01875455,0.00044147318,0.0008404264,0.0035647932,0.00023829895,0.0013120739,0.0012530104,0.0007234714,0.11370514],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003630374,0.00009297599,0.0027271956,0.0023424171,0.00005830915,0.0001381781,0.000075683085,0.0004647164,0.00059998134,0.0007379249,0.9737857,0.018613936],"study_design_scores_gemma":[0.0021158403,0.00020999773,0.024722127,0.0013534038,0.0002107716,0.0007118265,0.0003476576,0.00457367,0.005519207,0.008421266,0.95165807,0.00015606242],"about_ca_topic_score_codex":0.0036780739,"about_ca_topic_score_gemma":0.0059268624,"teacher_disagreement_score":0.6593007,"about_ca_system_score_codex":0.0009827464,"about_ca_system_score_gemma":0.001942426,"threshold_uncertainty_score":0.4859662},"labels":[],"label_agreement":null},{"id":"W6976993431","doi":"10.60692/b24zv-5v986","title":"HIV & Hepatitis in the Americas 28–30 April 2016, Mexico City, Mexico","year":2016,"lang":"en","type":"article","venue":"Greater South Information System","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Human immunodeficiency virus (HIV); Hepatitis C; Viral disease; Public health; Viral hepatitis","score_opus":0.027136831950922567,"score_gpt":0.23794942287853707,"score_spread":0.2108125909276145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6976993431","genre_codex":"empirical","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7322868,0.020045854,0.0005718243,0.10580616,0.003104588,0.00013434868,0.024903623,0.00015358122,0.11299331],"genre_scores_gemma":[0.8019205,0.015163105,0.0011405431,0.0062109362,0.0018845093,0.00017577564,0.013538197,0.000043063515,0.15992343],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99992955,0.0000104169285,0.0000037611653,0.0000124500375,0.000014208899,0.000029607423],"domain_scores_gemma":[0.9998275,0.000019280473,0.00006108448,0.000006286218,0.000030180006,0.000055585282],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.000288882,0.00010533271,0.00006858579,0.0004554309,0.0015456951,0.00078099384,0.00017954184,0.00041650896,0.008138016],"category_scores_gemma":[0.00044896925,0.00010177948,0.00010093273,0.0005107258,0.00032695464,0.00043047278,0.0008121263,0.00067573786,0.00048041882],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00026036627,0.00011618931,0.48870507,0.00032915262,0.000060373135,0.0018901916,0.007837353,0.0003567101,0.001174253,0.004577465,0.38928586,0.105407074],"study_design_scores_gemma":[0.00001361163,0.000017290788,0.7502022,0.00022976549,0.000017982367,0.00042759313,0.009025597,0.00010883236,0.0002231777,0.00034781726,0.2393779,0.000008239432],"about_ca_topic_score_codex":0.14278923,"about_ca_topic_score_gemma":0.3664591,"teacher_disagreement_score":0.991862,"about_ca_system_score_codex":0.002021019,"about_ca_system_score_gemma":0.0017097653,"threshold_uncertainty_score":0.28391623},"labels":[],"label_agreement":null},{"id":"W6977030600","doi":"10.6084/m9.figshare.12228167.v1","title":"Additional file 1 of The PRECISE (PREgnancy Care Integrating translational Science, Everywhere) database: open-access data collection in maternal and newborn health","year":2020,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of British Columbia","funders":"","keywords":"Data collection; Health care; Health data; MEDLINE; Patient data","score_opus":0.10922254217658944,"score_gpt":0.36451374350244403,"score_spread":0.2552912013258546,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977030600","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00005797201,0.000016423883,0.00015444068,0.000057064422,0.000010710252,0.000033758162,0.99902225,0.00016992446,0.00047746362],"genre_scores_gemma":[0.003230829,0.00017500544,0.0031316853,0.00040794347,0.000062587715,0.0010128552,0.98818797,0.000731538,0.0030596228],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99882394,0.00024535996,0.0002858807,0.0002870185,0.00022415271,0.00013355151],"domain_scores_gemma":[0.9728354,0.021302562,0.0014497491,0.001130948,0.0024399266,0.0008413621],"candidate_categories":["open_science","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0024601887,0.0010850421,0.0011608535,0.0039773537,0.0006977195,0.0020337687,0.0018540563,0.0012426404,0.7467743],"category_scores_gemma":[0.031076506,0.0006209772,0.00088959123,0.0070708725,0.00038849175,0.0019160341,0.0017866028,0.0011559635,0.15129432],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021620386,0.00004404184,0.001867338,0.0029521075,0.000046089346,0.000042456053,0.00006803546,0.00022228817,0.00009875368,0.00080270215,0.9865947,0.0070453747],"study_design_scores_gemma":[0.0016573055,0.000097411,0.015827935,0.0026518658,0.00018710629,0.00033036852,0.00044350716,0.0006730278,0.0008109371,0.009166613,0.968037,0.000116997035],"about_ca_topic_score_codex":0.007907546,"about_ca_topic_score_gemma":0.012705871,"teacher_disagreement_score":0.99814594,"about_ca_system_score_codex":0.0012536591,"about_ca_system_score_gemma":0.002438771,"threshold_uncertainty_score":0.36119568},"labels":[],"label_agreement":null},{"id":"W6977146106","doi":"10.6084/m9.figshare.20175959","title":"Additional file 11 of The Semanticscience Integrated Ontology (SIO) for biomedical research and knowledge discovery","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of New Brunswick; Ontario Institute for Cancer Research; Carleton University","funders":"","keywords":"Ontology; Knowledge extraction; File format; Key (lock); Flat file database","score_opus":0.07574097540738355,"score_gpt":0.34069025744204695,"score_spread":0.2649492820346634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977146106","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000120281176,0.000016746968,0.0009989525,0.00013739555,0.00005297823,0.000054611894,0.99426043,0.0018129918,0.002545533],"genre_scores_gemma":[0.0040266044,0.00016304046,0.010116097,0.00040682853,0.00007856628,0.0006957431,0.97046036,0.004391591,0.009661081],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994407,0.000087542816,0.00009265696,0.0001359261,0.0001508093,0.00009231764],"domain_scores_gemma":[0.9892603,0.007568349,0.0004861669,0.0007303727,0.0015805988,0.00037414],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016732567,0.0011551693,0.0010194086,0.0036622745,0.0008586453,0.0024855635,0.0018402536,0.0012402016,0.8182728],"category_scores_gemma":[0.016357902,0.00073136576,0.0013547912,0.0052901115,0.0004521415,0.0034810947,0.0018325826,0.0013894307,0.2743302],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016389645,0.0000470803,0.00060516153,0.001303879,0.00002824644,0.000056625115,0.000082601626,0.0002727967,0.00020488836,0.0019330092,0.9876916,0.0076100905],"study_design_scores_gemma":[0.0009716824,0.000035420544,0.0037255061,0.0007864697,0.00006592624,0.00020828516,0.00024446542,0.0010535556,0.0011480019,0.012013037,0.97966945,0.000078237244],"about_ca_topic_score_codex":0.008647431,"about_ca_topic_score_gemma":0.0115243355,"teacher_disagreement_score":0.8182728,"about_ca_system_score_codex":0.0016069448,"about_ca_system_score_gemma":0.0023406914,"threshold_uncertainty_score":0.25921166},"labels":[],"label_agreement":null},{"id":"W6977210199","doi":"10.6084/m9.figshare.21397896.v1","title":"Additional file 3 of Machine learning algorithms to identify cluster randomized trials from MEDLINE and EMBASE","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Western University; Ottawa Hospital; McMaster University; London Health Sciences Centre; Lawson Health Research Institute","funders":"","keywords":"Cluster (spacecraft); Convolutional neural network; MEDLINE; Feature (linguistics); Key (lock)","score_opus":0.045879010219779,"score_gpt":0.32122997232150885,"score_spread":0.2753509621017298,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977210199","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00016080111,0.00004037833,0.0004592693,0.000090554204,0.00001806008,0.00012133322,0.99813807,0.00050532224,0.00046610998],"genre_scores_gemma":[0.00969416,0.0003307338,0.0106572425,0.0006708838,0.00015186757,0.004743794,0.9642343,0.0015790577,0.007938044],"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99862695,0.0002583385,0.00038253065,0.0003468046,0.00024315451,0.00014222004],"domain_scores_gemma":[0.9507465,0.04223824,0.0023685636,0.0014561303,0.0026429915,0.00054766174],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0029725316,0.0014591111,0.0018794806,0.0047663064,0.00063985673,0.0020365752,0.0020083906,0.0015724734,0.8534177],"category_scores_gemma":[0.046422902,0.0008275633,0.0018407886,0.0072019347,0.00040035188,0.0021957303,0.0013407909,0.0010629107,0.12418414],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011180661,0.00015526045,0.003133337,0.02215332,0.00031767332,0.00017178425,0.00008893523,0.0015960104,0.00034494032,0.0018407514,0.9484676,0.020612322],"study_design_scores_gemma":[0.019630754,0.00082114746,0.028544622,0.010246895,0.0011368139,0.0010309053,0.00039540243,0.008552861,0.0026070816,0.034428675,0.8922527,0.00035226427],"about_ca_topic_score_codex":0.0052995365,"about_ca_topic_score_gemma":0.0135017885,"teacher_disagreement_score":0.99702746,"about_ca_system_score_codex":0.0017595406,"about_ca_system_score_gemma":0.0028626309,"threshold_uncertainty_score":0.20908189},"labels":[],"label_agreement":null},{"id":"W6977215450","doi":"10.6084/m9.figshare.c.5676403.v1","title":"The effect of rehabilitation protocol using mobile health in overweight and obese patients with knee osteoarthritis: a clinical trial","year":2021,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Randomized controlled trial; Rehabilitation; Overweight; WOMAC; mHealth; Osteoarthritis; Clinical trial; Activities of daily living","score_opus":0.020161997112257458,"score_gpt":0.3568855682054809,"score_spread":0.33672357109322343,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977215450","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8618111,0.029852567,0.0027661962,0.0025585939,0.0040257014,0.090361714,0.0016172499,0.0003129769,0.0066939173],"genre_scores_gemma":[0.874751,0.008551047,0.007026381,0.0023138637,0.0018650587,0.102087736,0.0006120198,0.000024595856,0.0027683552],"study_design_codex":"randomized_trial","study_design_gemma":"randomized_trial","domain_scores_codex":[0.99295664,0.00462232,0.001055156,0.00055088464,0.00045028576,0.00036466087],"domain_scores_gemma":[0.9952792,0.002062101,0.0013636682,0.00026945156,0.00034445163,0.0006811241],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0061468454,0.0016520409,0.004152949,0.0010531232,0.0009880131,0.00163437,0.0010992347,0.0038068064,0.009966773],"category_scores_gemma":[0.011749904,0.00088491867,0.0038339985,0.0009823472,0.0014825489,0.0015363934,0.00082399754,0.0027171336,0.0009067009],"study_design_candidate":"randomized_trial","study_design_consensus":"randomized_trial","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.9608818,0.011581893,0.00049645855,0.0053335344,0.0023689254,0.00004471763,0.000073323165,0.00018520489,0.0005469352,0.00014905407,0.0004441811,0.017893959],"study_design_scores_gemma":[0.96259654,0.03514885,0.00060415803,0.00025520584,0.0007503049,0.000011297889,0.000022913193,0.00013635572,0.00008779891,0.00006539586,0.00031213652,0.000008979622],"about_ca_topic_score_codex":0.0014291737,"about_ca_topic_score_gemma":0.0020332045,"teacher_disagreement_score":0.009966773,"about_ca_system_score_codex":0.0011782544,"about_ca_system_score_gemma":0.0025715246,"threshold_uncertainty_score":0.033342123},"labels":[],"label_agreement":null},{"id":"W6977259263","doi":"10.6084/m9.figshare.19759386","title":"Additional file 4 of Automated medical chart review for breast cancer outcomes research: a novel natural language processing extraction system","year":2022,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Prevention of Organ Failure; University of British Columbia","funders":"","keywords":"Breast cancer; Chart; Data extraction; Natural language; Gold standard (test)","score_opus":0.06701088399541694,"score_gpt":0.40579832678477173,"score_spread":0.3387874427893548,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977259263","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002842426,0.000015688798,0.000916143,0.000115992,0.000023992454,0.00009330971,0.99631846,0.0013911112,0.0008410373],"genre_scores_gemma":[0.006613003,0.00008620294,0.009743799,0.0003210041,0.00009925096,0.0011934789,0.9753749,0.0012583352,0.005309931],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990619,0.00014558822,0.0002459843,0.00024168154,0.00021774948,0.00008726944],"domain_scores_gemma":[0.98153806,0.012956213,0.0013382597,0.001044091,0.0027188822,0.0004045137],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0016714409,0.0010658685,0.0007905045,0.0034535138,0.00054820534,0.0014694058,0.0015326638,0.00092765805,0.68614],"category_scores_gemma":[0.019467695,0.00045523277,0.00090460235,0.003643862,0.00024601741,0.0013598424,0.0012760698,0.0007297028,0.11694498],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003986916,0.000097667056,0.0028830061,0.002525373,0.00006484831,0.00013900112,0.000074955664,0.00049432163,0.00058036624,0.00073143,0.97342366,0.018586656],"study_design_scores_gemma":[0.002386896,0.00022009267,0.025412096,0.0014987725,0.00022775604,0.0006717266,0.00035015531,0.00458695,0.00558989,0.008969552,0.9499213,0.00016479808],"about_ca_topic_score_codex":0.004001167,"about_ca_topic_score_gemma":0.006318447,"teacher_disagreement_score":0.68614,"about_ca_system_score_codex":0.001078075,"about_ca_system_score_gemma":0.002018982,"threshold_uncertainty_score":0.44768316},"labels":[],"label_agreement":null},{"id":"W6977427991","doi":"10.6084/m9.figshare.26585853","title":"Additional file 2 of The use of text-mining software to facilitate screening of literature on centredness in health care","year":2024,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trinity Western University; Providence Health Care","funders":"","keywords":"Health care; Software; MEDLINE; Patient care; Data collection; Public health","score_opus":0.06566845128082391,"score_gpt":0.2762970640310374,"score_spread":0.21062861275021347,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977427991","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0001455698,0.000032730266,0.00025198443,0.000111572204,0.000016973276,0.00016126818,0.9981584,0.00030803107,0.00081335404],"genre_scores_gemma":[0.006101336,0.00038300623,0.0092999935,0.00070575305,0.00011619793,0.005607231,0.96633565,0.0013045075,0.010146356],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9988048,0.00022267744,0.00041490485,0.00023760788,0.00019562061,0.0001243407],"domain_scores_gemma":[0.9521942,0.040341068,0.0022929457,0.0009872168,0.0033712653,0.00081340485],"candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0026138595,0.0012058263,0.0016641223,0.0073666796,0.00091591256,0.002146227,0.0018716458,0.0013544306,0.8444447],"category_scores_gemma":[0.03691411,0.00059228024,0.0015005613,0.008657824,0.00042740442,0.0024956332,0.0018133147,0.0010486171,0.12709253],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006952039,0.00009945361,0.0021513123,0.022093218,0.0001328099,0.00018012486,0.0002515657,0.0002993668,0.00030683167,0.0014666269,0.9567811,0.015542303],"study_design_scores_gemma":[0.0043508788,0.00023628105,0.017930066,0.00984006,0.0005640624,0.0007334484,0.00082631106,0.0011088796,0.0015064691,0.012746898,0.9499611,0.00019548264],"about_ca_topic_score_codex":0.0062452173,"about_ca_topic_score_gemma":0.01253244,"teacher_disagreement_score":0.99738616,"about_ca_system_score_codex":0.0015182333,"about_ca_system_score_gemma":0.0035152077,"threshold_uncertainty_score":0.2218808},"labels":[],"label_agreement":null},{"id":"W6977454423","doi":"10.6084/m9.figshare.c.3629726_d5","title":"Additional file 3: of Predicting potential ranges of primary malaria vectors and malaria in northern South America based on projected changes in climate, land cover and human population","year":2015,"lang":"en","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Malaria; Precipitation; Land cover; Population; Variable (mathematics); Quarter (Canadian coin); Bar (unit)","score_opus":0.020846880799051237,"score_gpt":0.23526658719688662,"score_spread":0.2144197063978354,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977454423","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0002883234,0.000008638667,0.00037476383,0.00010172506,0.000016370923,0.000021644972,0.9975079,0.0007872734,0.00089318],"genre_scores_gemma":[0.011663278,0.000078244346,0.0057553905,0.00019395992,0.00004155003,0.00048172154,0.9742192,0.001278316,0.006288337],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997495,0.000041238636,0.000037008864,0.00007852707,0.00006179221,0.0000319725],"domain_scores_gemma":[0.99535024,0.0034869905,0.0001706173,0.00028766243,0.0005770698,0.00012740219],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00082354306,0.0011480378,0.0007336344,0.0015885797,0.00054480415,0.0013876554,0.0015937724,0.0009876258,0.74079823],"category_scores_gemma":[0.01094561,0.000489132,0.0010092907,0.0025375397,0.00021381822,0.0016707638,0.00081192434,0.0007051797,0.14998414],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016784813,0.000060110706,0.0036472518,0.0014642691,0.00006650458,0.00009995523,0.00007820678,0.0023015249,0.00017853528,0.0006843265,0.9804504,0.010801134],"study_design_scores_gemma":[0.003229587,0.00018806968,0.044948805,0.0018340696,0.00027752836,0.0004588847,0.00079031073,0.021982444,0.0026425999,0.014901542,0.9085598,0.00018639641],"about_ca_topic_score_codex":0.017266257,"about_ca_topic_score_gemma":0.026095286,"teacher_disagreement_score":0.74079823,"about_ca_system_score_codex":0.0009599069,"about_ca_system_score_gemma":0.0013659553,"threshold_uncertainty_score":0.36971986},"labels":[],"label_agreement":null},{"id":"W6977790841","doi":"10.6084/m9.figshare.c.6962033.v1","title":"Patient-reported experiences and outcomes of virtual care during COVID-19: a systematic review","year":2024,"lang":"en","type":"other","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Systematic review; MEDLINE; Health care; Quality (philosophy); Critical appraisal; Virtual patient; Health professionals; Best practice","score_opus":0.02795715729872682,"score_gpt":0.3163951956076763,"score_spread":0.2884380383089495,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6977790841","genre_codex":"review","genre_gemma":"review","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"review","genre_consensus":"review","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0047453674,0.99314016,0.0001633032,0.00028347995,0.00007375936,0.0006005283,0.0006652247,0.00000627503,0.00032198452],"genre_scores_gemma":[0.07665416,0.91946125,0.0008896645,0.00058416644,0.00010271876,0.0016973417,0.00047967443,0.000007072038,0.00012392977],"study_design_codex":"systematic_review","study_design_gemma":"systematic_review","domain_scores_codex":[0.97983164,0.008300976,0.006692281,0.0013008404,0.003468625,0.00040566936],"domain_scores_gemma":[0.9305578,0.052367963,0.011991713,0.00084393605,0.0037850193,0.00045362982],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0141361095,0.0010768607,0.006393944,0.008485364,0.000739209,0.0036881242,0.0017017485,0.00160684,0.0035956213],"category_scores_gemma":[0.07348461,0.0009428032,0.008122983,0.01122393,0.0011875135,0.0027640946,0.0018814254,0.0012615662,0.00023102202],"study_design_candidate":"systematic_review","study_design_consensus":"systematic_review","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002816362,0.000029412364,0.0028256155,0.95332444,0.012924957,0.00011544468,0.0006126881,0.00008680323,0.00008105478,0.00017298826,0.0008916889,0.028653285],"study_design_scores_gemma":[0.00027763494,0.0004054346,0.013078065,0.90837145,0.06544958,0.000553712,0.00130849,0.00009815252,0.00020324918,0.00026351935,0.009926913,0.00006377813],"about_ca_topic_score_codex":0.005670317,"about_ca_topic_score_gemma":0.013995313,"teacher_disagreement_score":0.0141361095,"about_ca_system_score_codex":0.0044834656,"about_ca_system_score_gemma":0.010948104,"threshold_uncertainty_score":0.07475978},"labels":[],"label_agreement":null},{"id":"W6981108759","doi":"","title":"Die Einführung zur actio finium regundorum ―besonders in der Einstellung des Fokus auf die Kontroverse zwischen Focke Tannen Hinrichs und Rolf Knütel","year":2015,"lang":"de","type":"other","venue":"Meiji Gakuin University Institutional Repository (Meiji Gakuin University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Middle Ages; Medieval history; Quarter (Canadian coin)","score_opus":0.015266945349007606,"score_gpt":0.2192599071442296,"score_spread":0.203992961795222,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6981108759","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.035591424,0.020694355,0.062308073,0.21412788,0.011881859,0.00068016147,0.022424884,0.009036725,0.62325466],"genre_scores_gemma":[0.22451685,0.02165615,0.115347184,0.018039376,0.0022143908,0.001349259,0.030527705,0.0061489716,0.58020014],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99532145,0.0017587272,0.00037880754,0.00048549668,0.0016421948,0.00041328635],"domain_scores_gemma":[0.9917497,0.0029629823,0.000622441,0.0015665082,0.0020043245,0.0010940329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.013068534,0.0005877009,0.0004227458,0.0029144748,0.002417495,0.005309941,0.0014046644,0.0018674085,0.033655956],"category_scores_gemma":[0.018177744,0.0003354801,0.00047679828,0.003138381,0.001580748,0.0047228993,0.005998672,0.0034525162,0.016547248],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00015994704,0.00010259474,0.0026319004,0.0004406876,0.000025353133,0.00024897425,0.0024713406,0.0005451204,0.0019063904,0.10798027,0.46903357,0.4144539],"study_design_scores_gemma":[0.0000077730465,0.000010505143,0.0015341981,0.00016674581,0.0000067010474,0.00006914866,0.00032103903,0.00025325053,0.0010112688,0.0043695276,0.9922382,0.000011723839],"about_ca_topic_score_codex":0.01712054,"about_ca_topic_score_gemma":0.012043183,"teacher_disagreement_score":0.033655956,"about_ca_system_score_codex":0.0041693514,"about_ca_system_score_gemma":0.012971876,"threshold_uncertainty_score":0.11259043},"labels":[],"label_agreement":null},{"id":"W6981706237","doi":"","title":"An experimental analysis of chip formation in circular micro-end-milling","year":2010,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Chip; Chip formation; Fabrication; Automotive industry; Focus (optics); Die (integrated circuit)","score_opus":0.012888355024390332,"score_gpt":0.2838288543013835,"score_spread":0.27094049927699315,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6981706237","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9765453,0.0003205836,0.020413851,0.00004167383,0.000033224384,0.00007254382,0.00022047135,0.00031741487,0.002034919],"genre_scores_gemma":[0.9865693,0.00009020242,0.012292228,0.000024020726,0.0000045598827,0.00003352371,0.00012919199,0.000026197906,0.00083075627],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99945635,0.00004147079,0.000026672848,0.0001467852,0.00023208972,0.0000965595],"domain_scores_gemma":[0.9978934,0.0008213805,0.00030050168,0.0004281438,0.00046106856,0.00009556424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004975259,0.00021309916,0.00026091255,0.00028961833,0.00032948016,0.00030344547,0.00044862746,0.0006018144,0.0016826662],"category_scores_gemma":[0.0012471736,0.00024122796,0.00018957655,0.00038691933,0.00067518046,0.0003138713,0.0002855144,0.00034801452,0.00025824242],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00021116043,0.00010002814,0.0010786619,0.0000966467,0.00000866306,0.00017206017,0.0002730902,0.0022127386,0.98740965,0.00027713188,0.00015395864,0.008006068],"study_design_scores_gemma":[0.000014566025,0.0006367784,0.009425064,0.000006100138,0.000010632531,0.00013476903,0.000109376444,0.011013836,0.97735983,0.000099033576,0.0011722194,0.000017944014],"about_ca_topic_score_codex":0.00049583695,"about_ca_topic_score_gemma":0.0006477177,"teacher_disagreement_score":0.0016826662,"about_ca_system_score_codex":0.00030326183,"about_ca_system_score_gemma":0.00019806331,"threshold_uncertainty_score":0.0056290627},"labels":[],"label_agreement":null},{"id":"W6982361456","doi":"","title":"Imperial (Canadian 94887)","year":2011,"lang":"en","type":"other","venue":"OhioLink ETD Center (Ohio Library and Information Network)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Bay; Georgian; Hull; Shore","score_opus":0.008154667103709445,"score_gpt":0.1987039319811406,"score_spread":0.19054926487743118,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6982361456","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00066122133,0.00047356394,0.00019391377,0.0005983666,0.00021244377,0.000021298632,0.011638501,0.00040884153,0.9857918],"genre_scores_gemma":[0.0014243649,0.00017900892,0.00011677864,0.00014395393,0.0000141507335,0.0000048870616,0.0026426616,0.0001506804,0.9953235],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99953365,0.000019177081,0.00001568161,0.000094666604,0.00020380742,0.00013303552],"domain_scores_gemma":[0.99910223,0.000035720903,0.000024820947,0.000060736358,0.0005085931,0.00026794916],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00032488527,0.000879947,0.0005615229,0.0019696981,0.0058584968,0.0045106728,0.000803565,0.0011993793,0.75914246],"category_scores_gemma":[0.0012542785,0.00042583526,0.00035901647,0.0040848735,0.0006615073,0.0016327502,0.001751226,0.0011158878,0.51288146],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034307093,0.000011429466,0.00059086265,0.000059181766,0.0000025561048,0.00016536209,0.00013516154,0.00005432456,0.00019604675,0.006047912,0.9424422,0.050260734],"study_design_scores_gemma":[0.0000015999311,0.0000018171694,0.00077355455,0.000018884484,0.0000010657677,0.000042771793,0.0000811136,0.000014248585,0.00005071411,0.00011191627,0.99889886,0.0000035570035],"about_ca_topic_score_codex":0.72910154,"about_ca_topic_score_gemma":0.934994,"teacher_disagreement_score":0.72910154,"about_ca_system_score_codex":0.009822196,"about_ca_system_score_gemma":0.011407272,"threshold_uncertainty_score":0.5449877},"labels":[],"label_agreement":null},{"id":"W6984340161","doi":"","title":"IB at St. Robert","year":2006,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.004636347176740512,"score_gpt":0.18957071695718306,"score_spread":0.18493436978044253,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6984340161","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0013883598,0.00095185626,0.0002519329,0.003133556,0.001681148,0.00006140575,0.0026953013,0.00094496016,0.9888915],"genre_scores_gemma":[0.002708082,0.00027459068,0.00019901576,0.00033609045,0.00008867755,0.000012713471,0.00064427784,0.0001776882,0.99555886],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996232,0.000029778366,0.000010549904,0.000093098366,0.00013634446,0.00010699136],"domain_scores_gemma":[0.9993517,0.000039528262,0.000029230072,0.000048163387,0.00022598206,0.00030543664],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00033403965,0.0008484009,0.00040856836,0.00093744544,0.0041518174,0.003408322,0.0005355962,0.0013643963,0.7673485],"category_scores_gemma":[0.0011569593,0.00034840108,0.00038888215,0.00075242197,0.00037964684,0.0011751956,0.001998927,0.0016800746,0.50314003],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007697739,0.000037760477,0.00055203115,0.00010159349,0.00000395277,0.00028727055,0.00017415002,0.000043428896,0.00073310104,0.0038705096,0.94670284,0.047416426],"study_design_scores_gemma":[0.0000042587044,0.000010108034,0.001057842,0.000046880366,0.0000019175031,0.000073817624,0.00013562903,0.000017671075,0.00013534528,0.0002113107,0.99830127,0.0000038654107],"about_ca_topic_score_codex":0.047756035,"about_ca_topic_score_gemma":0.18843634,"teacher_disagreement_score":0.23265147,"about_ca_system_score_codex":0.0028547093,"about_ca_system_score_gemma":0.0033060813,"threshold_uncertainty_score":0.33184904},"labels":[],"label_agreement":null},{"id":"W6986366228","doi":"","title":"Pelvic osteosarcoma ressection in a bitch: case report","year":2009,"lang":"en","type":"article","venue":"LA Referencia (Red Federada de Repositorios Institucionales de Publicaciones Científicas)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Osteosarcoma; Carboplatin; Pelvis; Chemotherapy; Urinary system; Quality of life (healthcare)","score_opus":0.020491986376615538,"score_gpt":0.2659861284623429,"score_spread":0.24549414208572734,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6986366228","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9902035,0.0023298047,0.0015304856,0.0008933597,0.00015032878,0.00016624223,0.00023844825,0.00007508554,0.004412616],"genre_scores_gemma":[0.99536514,0.0010609811,0.0010405268,0.0004062923,0.00021586893,0.000021990025,0.00010677729,0.000015409038,0.0017670444],"study_design_codex":"case_report","study_design_gemma":"case_report","domain_scores_codex":[0.99942684,0.000046110905,0.00005742646,0.00019468406,0.00007398618,0.00020094754],"domain_scores_gemma":[0.9991647,0.00020320536,0.000249621,0.0000796567,0.00004787743,0.00025497776],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002567062,0.0018479009,0.0011937022,0.0020883589,0.0026461945,0.0015944313,0.00086492306,0.005440815,0.0037873425],"category_scores_gemma":[0.0018003448,0.0012603777,0.0010353678,0.0010655202,0.0017170748,0.0012548658,0.001496094,0.0020061724,0.0013815962],"study_design_candidate":"case_report","study_design_consensus":"case_report","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026830596,0.000030617284,0.0033527024,0.00002588073,0.0000045487577,0.99518317,0.0001396942,0.000027347727,0.00059679383,0.00004817184,0.000051307907,0.00051293656],"study_design_scores_gemma":[0.000009903507,0.000116156945,0.006278334,0.000015867803,0.00001615428,0.9921835,0.00027328962,0.00013181777,0.00051235646,0.00010097753,0.00035117764,0.000010438902],"about_ca_topic_score_codex":0.0021289021,"about_ca_topic_score_gemma":0.0043472312,"teacher_disagreement_score":0.005440815,"about_ca_system_score_codex":0.001074794,"about_ca_system_score_gemma":0.0005454968,"threshold_uncertainty_score":0.012669921},"labels":[],"label_agreement":null},{"id":"W6986959462","doi":"","title":"Royal Bank : Bay and Temperanke : Holdup","year":2013,"lang":"en","type":"other","venue":"York University Digital Library (York University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Bay; Hydrology (agriculture); Shore; Port (circuit theory)","score_opus":0.006312922088334687,"score_gpt":0.1537580356362193,"score_spread":0.1474451135478846,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6986959462","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.011731437,0.0016204759,0.0017747168,0.008002366,0.0010284946,0.0000675827,0.15013598,0.007977575,0.8176613],"genre_scores_gemma":[0.027189266,0.0010371343,0.0018631397,0.0006047809,0.00015281561,0.00002353993,0.03298683,0.0022494062,0.933893],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9997732,0.000011150396,0.000005612533,0.00002843126,0.0001328503,0.000048741033],"domain_scores_gemma":[0.9990901,0.00007493864,0.000049311402,0.00004865341,0.00041272576,0.00032430424],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00032932972,0.00049775455,0.0003106336,0.0018491317,0.0033160832,0.004998956,0.0006623716,0.00086817663,0.2701207],"category_scores_gemma":[0.0015415486,0.00021923815,0.00020003597,0.003657015,0.0007273011,0.0018533478,0.0015054912,0.0009106494,0.06843686],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003989221,0.000005746141,0.0010744763,0.000059049995,0.0000024006445,0.00015629386,0.00060803513,0.00004774591,0.00015321371,0.0016446102,0.9751153,0.021093242],"study_design_scores_gemma":[0.0000027057715,0.0000021886233,0.00568903,0.000043203916,0.0000025074849,0.000083074905,0.0012056754,0.000089194386,0.00021773981,0.00023475265,0.9924189,0.0000110310275],"about_ca_topic_score_codex":0.60244334,"about_ca_topic_score_gemma":0.84830105,"teacher_disagreement_score":0.39755666,"about_ca_system_score_codex":0.00454838,"about_ca_system_score_gemma":0.0047095926,"threshold_uncertainty_score":0.9036438},"labels":[],"label_agreement":null},{"id":"W6991712822","doi":"","title":"IFSM, wavelets and fractal-wavelets, three methods of approximation","year":2006,"lang":"en","type":"dissertation","venue":"Library and Archives Canada (Government of Canada)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Wavelet; Generalization; Noise (video); Representation (politics); Feature (linguistics); Stability (learning theory)","score_opus":0.004571880292447613,"score_gpt":0.19791290422243507,"score_spread":0.19334102392998745,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6991712822","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.003864747,0.0074111447,0.97999316,0.0012222589,0.0005278305,0.000033001867,0.00017694246,0.00034215164,0.0064287353],"genre_scores_gemma":[0.13591737,0.013999732,0.8297611,0.000492804,0.0015784517,0.0002625287,0.00076439884,0.0003840541,0.01683966],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9983974,0.00044253378,0.0001148609,0.0001466803,0.0007999954,0.00009850138],"domain_scores_gemma":[0.9974376,0.0011973695,0.0002144854,0.00041463258,0.000644122,0.00009177701],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003003583,0.0009414205,0.0011174999,0.0031204415,0.0007386541,0.0043721185,0.0012029313,0.0017192119,0.0034710981],"category_scores_gemma":[0.010906302,0.00031143008,0.001143515,0.0044947746,0.0029598258,0.0029885834,0.001312132,0.0022603094,0.0015157812],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001360851,0.000032128446,0.0013528877,0.00036929589,0.000062105806,0.00007789085,0.00035671509,0.023027435,0.003260525,0.60160244,0.0124264015,0.35729605],"study_design_scores_gemma":[0.000045807963,0.00007341065,0.0026240423,0.0002067157,0.00008961557,0.000497325,0.00034972315,0.2762713,0.0031945205,0.65994,0.056618467,0.00008908963],"about_ca_topic_score_codex":0.006728491,"about_ca_topic_score_gemma":0.0057866676,"teacher_disagreement_score":0.006728491,"about_ca_system_score_codex":0.0017343509,"about_ca_system_score_gemma":0.0019426451,"threshold_uncertainty_score":0.015884638},"labels":[],"label_agreement":null},{"id":"W6992361749","doi":"","title":"Le chantier épique d’Hugues Salel : La construction de la première version métrique de l’<i>Iliade</i> en français","year":2019,"lang":"fr","type":"article","venue":"Project Muse (Johns Hopkins University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Context (archaeology); Closure (psychology); Volume (thermodynamics); Work (physics)","score_opus":0.005684268831405992,"score_gpt":0.21378632231424718,"score_spread":0.2081020534828412,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6992361749","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.04720931,0.23385592,0.08195577,0.27327743,0.08744767,0.00060904655,0.021283729,0.002252964,0.25210822],"genre_scores_gemma":[0.32835364,0.17261603,0.10810557,0.03553533,0.016172068,0.0014491087,0.020239163,0.0064572357,0.31107187],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99290943,0.0021231442,0.00086593546,0.0008834634,0.002725404,0.000492617],"domain_scores_gemma":[0.9836932,0.004949844,0.00072984654,0.0008288223,0.00907822,0.0007200267],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.011361714,0.0011938751,0.00093146064,0.011796983,0.0050328914,0.010239319,0.0015134312,0.0022928885,0.017504053],"category_scores_gemma":[0.024125414,0.0005793317,0.0009546221,0.011613431,0.0053312066,0.0067248046,0.0032766808,0.004116678,0.00467931],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00025069245,0.00004110638,0.0041156868,0.0024123013,0.00011047812,0.0006470498,0.031978533,0.0005553221,0.003210963,0.220484,0.50652456,0.22966932],"study_design_scores_gemma":[0.000008191484,0.000008134321,0.0017070656,0.00080260064,0.000017632847,0.00013505146,0.0020507833,0.000084897656,0.00041227337,0.001154734,0.99358565,0.0000330239],"about_ca_topic_score_codex":0.45990625,"about_ca_topic_score_gemma":0.33255902,"teacher_disagreement_score":0.45990625,"about_ca_system_score_codex":0.016379293,"about_ca_system_score_gemma":0.026578873,"threshold_uncertainty_score":0.9144586},"labels":[],"label_agreement":null},{"id":"W6992993283","doi":"","title":"Multi-objective Representation for Numbers in Clinical Narratives Using CamemBERT-bio","year":2024,"lang":"en","type":"preprint","venue":"Open Repository and Bibliography (University of Luxembourg)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Fonds de Recherche du Québec - Santé; Institut de Valorisation des Données; Université de Montréal","keywords":"Representation (politics); Key (lock); Embedding; Simple (philosophy); Component (thermodynamics); Narrative","score_opus":0.06711852228574952,"score_gpt":0.36938067668208796,"score_spread":0.3022621543963384,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6992993283","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.056837294,0.0023197103,0.892192,0.0021004234,0.00043181595,0.00054022827,0.02053507,0.015007778,0.010035638],"genre_scores_gemma":[0.32883015,0.0011730077,0.63060975,0.0005355473,0.00021828253,0.000767688,0.031532153,0.00043561388,0.0058977497],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99921,0.00020858533,0.00010042079,0.00025081733,0.00019242482,0.00003778291],"domain_scores_gemma":[0.99850625,0.0007956105,0.00023959906,0.00016364545,0.00023335205,0.000061539526],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0011739541,0.0009676011,0.0004009683,0.0040522106,0.00033443287,0.0019197646,0.0008955811,0.0010481921,0.005473343],"category_scores_gemma":[0.006833405,0.00018735709,0.00091409177,0.001856814,0.00041632718,0.0026186362,0.0013074547,0.00075470406,0.0022223177],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087654544,0.0003008426,0.010760973,0.0016794741,0.00020957751,0.001127879,0.0012627214,0.06917524,0.019288877,0.061098807,0.043684807,0.7905343],"study_design_scores_gemma":[0.000078819496,0.00022662713,0.0072261617,0.0003350137,0.00009551596,0.0007399663,0.0006571778,0.8101135,0.014239541,0.06790996,0.09828636,0.00009132128],"about_ca_topic_score_codex":0.0037486723,"about_ca_topic_score_gemma":0.0048675234,"teacher_disagreement_score":0.005473343,"about_ca_system_score_codex":0.0014235075,"about_ca_system_score_gemma":0.0009101675,"threshold_uncertainty_score":0.01831019},"labels":[],"label_agreement":null},{"id":"W6997621323","doi":"","title":"Why choose this one? Factors in scientists' selection of bioinformatics tools","year":2011,"lang":"en","type":"article","venue":"eScholarship@McGill (McGill)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"Dalhousie University","keywords":"Selection (genetic algorithm); Identification (biology); Variation (astronomy); Usability; Feature selection; Sequence (biology)","score_opus":0.04235973405650275,"score_gpt":0.24716612020841922,"score_spread":0.20480638615191646,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W6997621323","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.93115264,0.0022478988,0.020871382,0.024074255,0.00033024597,0.00047263614,0.0002644337,0.00030115497,0.020285456],"genre_scores_gemma":[0.98510724,0.00073122414,0.010535458,0.0022300286,0.00010742543,0.00016907486,0.00011053865,0.000091330025,0.0009176618],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9097527,0.05468563,0.0079827905,0.004339328,0.01899954,0.0042399964],"domain_scores_gemma":[0.6683867,0.25075474,0.030602844,0.0057381797,0.03243802,0.012079532],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.066767804,0.00049181865,0.00062796084,0.00466946,0.005354487,0.008554213,0.00170287,0.0028545342,0.0023930967],"category_scores_gemma":[0.2526307,0.0006532864,0.0008950961,0.0039953035,0.005537501,0.006286335,0.0029124124,0.0030643255,0.0008790524],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009833042,0.0006089884,0.48011172,0.0025371795,0.0006020998,0.002270683,0.31694838,0.0011334958,0.006064803,0.008399105,0.019989012,0.16035114],"study_design_scores_gemma":[0.0003763822,0.0010827848,0.31835344,0.0022250498,0.0005825228,0.0039556907,0.5354991,0.0061096917,0.0053030336,0.027967198,0.097722135,0.0008229917],"about_ca_topic_score_codex":0.005018802,"about_ca_topic_score_gemma":0.005750628,"teacher_disagreement_score":0.9332322,"about_ca_system_score_codex":0.0036241312,"about_ca_system_score_gemma":0.0061413622,"threshold_uncertainty_score":0.35310614},"labels":[],"label_agreement":null},{"id":"W7002085825","doi":"","title":"Literature Mining in Molecular Biology","year":2002,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Biomedical text mining; Reading (process); Domain (mathematical analysis); Process (computing); MEDLINE; Information extraction","score_opus":0.01469741343752993,"score_gpt":0.2649923506805892,"score_spread":0.25029493724305923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7002085825","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0077108652,0.19740024,0.68059707,0.014494132,0.0044608507,0.0023870703,0.016537836,0.008903245,0.06750871],"genre_scores_gemma":[0.038287368,0.11184521,0.7974522,0.0051133493,0.0032023292,0.002360536,0.02275572,0.00072028337,0.018262949],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.99056375,0.0032762878,0.00169826,0.0018202745,0.0024008236,0.00024062581],"domain_scores_gemma":[0.98061407,0.01311442,0.0018543261,0.0020457453,0.001865161,0.000506349],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008048507,0.0015295933,0.0027163015,0.02489823,0.0024726773,0.0086131,0.0037432788,0.0025014482,0.015813552],"category_scores_gemma":[0.023913566,0.0010536264,0.0024147755,0.028300984,0.0026950303,0.0090636285,0.0042510107,0.002836495,0.013435338],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013980428,0.00017818816,0.001730064,0.013887318,0.00051009253,0.0010787646,0.0010935802,0.0030731496,0.004030034,0.14932849,0.07667759,0.7482729],"study_design_scores_gemma":[0.000058596193,0.000086424574,0.0023488323,0.004837641,0.0002078555,0.0018908024,0.0005556055,0.0047167726,0.0032261906,0.22011483,0.76186633,0.00009003229],"about_ca_topic_score_codex":0.0013027844,"about_ca_topic_score_gemma":0.0014573836,"teacher_disagreement_score":0.02489823,"about_ca_system_score_codex":0.0018803574,"about_ca_system_score_gemma":0.0045186984,"threshold_uncertainty_score":0.052901626},"labels":[],"label_agreement":null},{"id":"W7009779543","doi":"","title":"Forging a new model for cooperation: the Health Science Information Consortium of Toronto.","year":2015,"lang":"en","type":"article","venue":"QSpace (Queen's University Library)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Forging; Information system; Information technology; Health information; Information science","score_opus":0.019854763896363602,"score_gpt":0.24246565495466404,"score_spread":0.22261089105830045,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7009779543","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.014904837,0.019041417,0.13943762,0.7215046,0.0025771097,0.00033243344,0.0066009876,0.0038284336,0.09177265],"genre_scores_gemma":[0.44261134,0.0234583,0.33246168,0.024856461,0.0017458964,0.00081852596,0.018799284,0.0029680962,0.15228038],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9880084,0.005341293,0.0008179701,0.0012960837,0.0034278778,0.0011083909],"domain_scores_gemma":[0.9608466,0.013578566,0.0012795747,0.006172855,0.008196381,0.00992607],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.025867097,0.00065363647,0.0006711633,0.0023266156,0.0059934794,0.014292732,0.0023765203,0.0045726877,0.013519907],"category_scores_gemma":[0.031895753,0.000986455,0.0005517645,0.0069631627,0.008357648,0.021278497,0.009830204,0.0033004265,0.003035251],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019916608,0.000040430663,0.002907767,0.00038088005,0.000040829764,0.0003518373,0.00843184,0.0017563645,0.00077764544,0.38985673,0.4847999,0.11045661],"study_design_scores_gemma":[0.0000515489,0.000027974367,0.0032634689,0.0004467144,0.000038508875,0.000118013486,0.0055981055,0.0045457445,0.00071514084,0.11807887,0.8670376,0.00007836766],"about_ca_topic_score_codex":0.40606532,"about_ca_topic_score_gemma":0.47824278,"teacher_disagreement_score":0.40606532,"about_ca_system_score_codex":0.022992244,"about_ca_system_score_gemma":0.057603996,"threshold_uncertainty_score":0.8074035},"labels":[],"label_agreement":null},{"id":"W7011714116","doi":"","title":"\\n40 Years of Pulsars : Millisecond Pulsars, Magnetars and More. McGill University, Montreal, Canada, 12-17 August 2007","year":2008,"lang":"en","type":"other","venue":"Radboud Repository (Radboud University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Magnetar; Millisecond pulsar; Pulsar; Neutron star; Pulsar planet","score_opus":0.0055445206049210185,"score_gpt":0.1734981648402996,"score_spread":0.16795364423537856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7011714116","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006688665,0.0019162194,0.0012057425,0.0017372613,0.00069938984,0.000057858582,0.92622715,0.003756783,0.063730724],"genre_scores_gemma":[0.00430216,0.0032413902,0.004922183,0.0005540774,0.0003816069,0.00008561768,0.7947947,0.0027382248,0.18898007],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99949384,0.00002322017,0.000031296997,0.00008446734,0.00028830045,0.00007881741],"domain_scores_gemma":[0.99784625,0.00020367379,0.00020614914,0.0002510502,0.0010008201,0.00049206003],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0010915264,0.0012398332,0.0009373445,0.0061922856,0.001829932,0.00405198,0.0019217639,0.00097579753,0.22103283],"category_scores_gemma":[0.0033471258,0.00066554913,0.0006100043,0.013334649,0.00060504087,0.002071024,0.0023924946,0.0009165547,0.10384239],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022469898,0.0000038544067,0.0004984333,0.000111307876,0.0000039597253,0.000016485264,0.000032608004,0.000045406592,0.0001155471,0.000542223,0.98147595,0.017131776],"study_design_scores_gemma":[0.000011051251,0.000002310729,0.0069030225,0.00008389484,0.0000061748156,0.000024353212,0.000045801444,0.00007796179,0.00017819222,0.0005926291,0.99206454,0.000010099455],"about_ca_topic_score_codex":0.48125172,"about_ca_topic_score_gemma":0.7710074,"teacher_disagreement_score":0.5187483,"about_ca_system_score_codex":0.0031671212,"about_ca_system_score_gemma":0.010241548,"threshold_uncertainty_score":0.956901},"labels":[],"label_agreement":null},{"id":"W7015812949","doi":"","title":"Undergraduate business education and pro-environmental behaviours: a multi group study","year":2008,"lang":"en","type":"dissertation","venue":"The Atrium (University of Guelph)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Bachelor; Business ethics; Business education; Higher education; Orientation (vector space); Undergraduate education","score_opus":0.01519541215540867,"score_gpt":0.24362991853701027,"score_spread":0.2284345063816016,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015812949","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99969184,0.00002244304,0.000026358253,0.000045334764,0.000002788763,0.000020910775,0.000009677001,9.2675634e-7,0.00017981412],"genre_scores_gemma":[0.99904734,0.000062423984,0.00010025403,0.00011363824,0.0000065964723,0.00006555148,0.000024371839,0.0000016611948,0.0005782067],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9981007,0.0006059189,0.00010773448,0.00020834027,0.0004726477,0.0005048238],"domain_scores_gemma":[0.99516463,0.0012648157,0.00092847407,0.000239312,0.0005223008,0.0018804399],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028956102,0.00028998108,0.0007471278,0.0014905476,0.0023959647,0.0020912834,0.00048469935,0.00091824296,0.002556283],"category_scores_gemma":[0.005406864,0.0004740445,0.0003169753,0.0010368732,0.00089954067,0.0013331058,0.0015398593,0.0012087327,0.0006149503],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00039310285,0.0039007233,0.8762433,0.000059237238,0.000062514504,0.00040720575,0.096314594,0.0000351888,0.0016050454,0.0001727787,0.00046824443,0.020338034],"study_design_scores_gemma":[0.00003699255,0.0024843337,0.8916946,0.000040104904,0.000024579449,0.00030760127,0.10311074,0.00018683597,0.0003186378,0.00011406424,0.001657383,0.000024098317],"about_ca_topic_score_codex":0.003168278,"about_ca_topic_score_gemma":0.0052320496,"teacher_disagreement_score":0.003168278,"about_ca_system_score_codex":0.0008666339,"about_ca_system_score_gemma":0.0011156013,"threshold_uncertainty_score":0.015313625},"labels":[],"label_agreement":null},{"id":"W7015925641","doi":"","title":"Unusual kinetic behaviour in cathodic H2 evolution from KF.2HF melts at mild-steel and alloy electrodes","year":2004,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Tafel equation; Anode; Cathodic protection; Electrolyte; Electrode; Polarization (electrochemistry)","score_opus":0.010447606902898761,"score_gpt":0.25096166640586404,"score_spread":0.24051405950296528,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7015925641","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99239737,0.0014559553,0.0032635403,0.00017460885,0.00001720764,0.000012444238,0.00028772102,0.00011119962,0.002280066],"genre_scores_gemma":[0.9979988,0.00026377547,0.0007063401,0.000017532691,0.000008053466,0.000005572715,0.00018975185,0.0000147902465,0.00079529145],"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99973756,0.000026555157,0.000017592407,0.00005367144,0.00011402059,0.000050544346],"domain_scores_gemma":[0.9997892,0.00009290643,0.000036371443,0.00001866227,0.000042608746,0.000020212621],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00021696897,0.00021386573,0.00024503891,0.00046447018,0.00026354846,0.00046036977,0.0003030335,0.0006258505,0.0012927325],"category_scores_gemma":[0.0005286265,0.00017120581,0.00019510108,0.00028834378,0.0002781333,0.0005253474,0.00020365967,0.0004628173,0.00021896156],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019783451,0.000014690201,0.001025004,0.00009197115,0.000017972257,0.00030631185,0.00015325229,0.00015452705,0.9928544,0.00020385571,0.00008967749,0.0048903823],"study_design_scores_gemma":[0.000004604151,0.000089568915,0.0051177093,0.0000033571928,0.0000073667643,0.000361897,0.00006170452,0.0009868285,0.9923515,0.000080060585,0.0009279462,0.000007538845],"about_ca_topic_score_codex":0.00090603554,"about_ca_topic_score_gemma":0.0010041293,"teacher_disagreement_score":0.0012927325,"about_ca_system_score_codex":0.00026918386,"about_ca_system_score_gemma":0.00008377805,"threshold_uncertainty_score":0.004324615},"labels":[],"label_agreement":null},{"id":"W7017005465","doi":"","title":"2018 Active Healthy Kids Scotland Report Card","year":2018,"lang":"en","type":"other","venue":"Strathprints: The University of Strathclyde institutional repository (University of Strathclyde)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Report card; Alliance; Public health; Physical activity; Smart card","score_opus":0.012197913174653548,"score_gpt":0.2189132077377433,"score_spread":0.20671529456308976,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017005465","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0017538745,0.0007425055,0.0016094072,0.008619969,0.00273549,0.00058910216,0.61957186,0.004655138,0.3597226],"genre_scores_gemma":[0.007520288,0.0014953754,0.004090541,0.0031104241,0.00115015,0.001495565,0.55526626,0.0031896487,0.42268175],"study_design_codex":"not_applicable","study_design_gemma":"observational","domain_scores_codex":[0.9943289,0.0006739608,0.0007655268,0.0004648019,0.0031390996,0.0006277504],"domain_scores_gemma":[0.97369385,0.003096726,0.0016025471,0.0016375486,0.016762316,0.0032070773],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0065166247,0.00087892934,0.0007723115,0.008254764,0.0017450001,0.006446566,0.0022524635,0.0016418908,0.32391134],"category_scores_gemma":[0.033062283,0.0005608153,0.00057644385,0.009115877,0.0006704802,0.0051827948,0.004226238,0.0011933577,0.23264894],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000021355592,0.000009248449,0.00029056333,0.0000836725,0.0000012646648,0.000028563125,0.000066385226,0.000009936038,0.000023413295,0.001502383,0.98995167,0.008011538],"study_design_scores_gemma":[0.000011034302,0.000002876713,0.0008300661,0.0000534684,0.0000012277992,0.0000145678205,0.00008400712,0.000014677725,0.000034076453,0.00021187901,0.9987357,0.0000063345515],"about_ca_topic_score_codex":0.070034236,"about_ca_topic_score_gemma":0.07000778,"teacher_disagreement_score":0.32391134,"about_ca_system_score_codex":0.0043245894,"about_ca_system_score_gemma":0.010769343,"threshold_uncertainty_score":0.96435845},"labels":[],"label_agreement":null},{"id":"W7017044672","doi":"","title":"Achievement inequity, students’ participation and curriculum offerings in mathematics and science Higher School Certificate (HSC) courses within a segregated school system in New South Wales, Australia","year":2020,"lang":"en","type":"dissertation","venue":"The Sydney eScholarship Repository (The University of Sydney)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Curriculum; Certificate; Government (linguistics); Socioeconomic status; Quarter (Canadian coin); Academic achievement; School system; Higher education","score_opus":0.045423996831111965,"score_gpt":0.28699967148278666,"score_spread":0.24157567465167468,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017044672","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9989643,0.000056107983,0.000025315669,0.00014214468,0.0000027308738,0.000009239493,0.00028612817,0.0000013193156,0.0005127611],"genre_scores_gemma":[0.9990804,0.00006312175,0.00005395733,0.000029656232,0.000002150502,0.000015579255,0.00022134362,0.0000013868831,0.0005324011],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9987931,0.00023712293,0.000110646586,0.00029302982,0.0002626635,0.00030340318],"domain_scores_gemma":[0.99752873,0.00018756694,0.0008946287,0.00010976749,0.00043939208,0.0008399468],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0012236194,0.00019943378,0.00030434623,0.0010382838,0.001065636,0.0011816123,0.00067345874,0.00037284647,0.0015442468],"category_scores_gemma":[0.0027030176,0.00030724617,0.00043585364,0.0012731925,0.00076442456,0.00093397073,0.0026779466,0.0011957422,0.0003340311],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000031419604,0.000110850335,0.98854744,0.000030613486,0.000034061355,0.00008184441,0.0070283017,0.00008195095,0.00014435772,0.00026632674,0.0004752087,0.003167458],"study_design_scores_gemma":[7.1015126e-7,0.000028814266,0.9971084,0.000012975809,0.0000038665767,0.000019768413,0.0024580134,0.00011700546,0.000018032006,0.0000245643,0.00020440726,0.0000034915888],"about_ca_topic_score_codex":0.41427025,"about_ca_topic_score_gemma":0.48937902,"teacher_disagreement_score":0.41427025,"about_ca_system_score_codex":0.002443939,"about_ca_system_score_gemma":0.0027507793,"threshold_uncertainty_score":0.82371783},"labels":[],"label_agreement":null},{"id":"W7017231177","doi":"","title":"Adaptação cultural, validade e confiabilidade do Inventário de Habilidades de Vida Independente – versão do paciente (ILSS-BR/P) com portadores de esquizofrenia","year":2015,"lang":"en","type":"other","venue":"LA Referencia (Red Federada de Repositorios Institucionales de Publicaciones Científicas)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Conselho Nacional de Desenvolvimento Científico e Tecnológico; Coordenação de Aperfeiçoamento de Pessoal de Nível Superior; Fundação de Amparo à Pesquisa do Estado de São Paulo","keywords":"Quality of life (healthcare); Reliability (semiconductor); Schizophrenia (object-oriented programming); Recall","score_opus":0.026615964016531553,"score_gpt":0.27277666052672694,"score_spread":0.2461606965101954,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017231177","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99411535,0.00072243623,0.0022115214,0.00020867286,0.00007075519,0.00032718177,0.00036811968,0.000021729742,0.0019541865],"genre_scores_gemma":[0.99713194,0.00026412832,0.0017843923,0.00003821899,0.000016690654,0.00027010203,0.00032930073,0.000010465653,0.0001547413],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9879822,0.005989923,0.0017700995,0.0008433494,0.0029460103,0.00046844097],"domain_scores_gemma":[0.9740423,0.013218521,0.0059617744,0.0021855053,0.0039027897,0.0006891225],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.022939453,0.0006783559,0.0008229292,0.002373333,0.000725032,0.0015739443,0.00083074614,0.0005808145,0.0008898315],"category_scores_gemma":[0.039626487,0.0004893699,0.0017784211,0.0016382976,0.0013842831,0.0009500975,0.0017123986,0.00090506655,0.00015181501],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011986817,0.00009614453,0.97746104,0.0001197272,0.0002905183,0.00007840618,0.0032314865,0.00027003748,0.00034725456,0.00019481833,0.00021404753,0.017576586],"study_design_scores_gemma":[0.000043826796,0.00034358804,0.9910471,0.00024000736,0.00018063432,0.0004211437,0.0033810541,0.0018132798,0.0005893259,0.0005999675,0.0013063872,0.000033622116],"about_ca_topic_score_codex":0.0038846536,"about_ca_topic_score_gemma":0.0061995345,"teacher_disagreement_score":0.022939453,"about_ca_system_score_codex":0.00091936707,"about_ca_system_score_gemma":0.0016112154,"threshold_uncertainty_score":0.12131691},"labels":[],"label_agreement":null},{"id":"W7017294599","doi":"","title":"An analysis of Canadian young adults’ eating behaviours towards sustainable food choices","year":2023,"lang":"en","type":"dissertation","venue":"UWSpace (University of Waterloo)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada","keywords":"Sustainability; Food choice; Consumption (sociology); Food systems; Eating behavior; Healthy eating; Climate change; Food consumption","score_opus":0.010079286766545003,"score_gpt":0.22794834263377128,"score_spread":0.21786905586722627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017294599","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99289167,0.00037440256,0.00005597012,0.00031889882,0.00001111188,0.000109238266,0.0025163589,0.000004451515,0.0037178402],"genre_scores_gemma":[0.991647,0.0014835241,0.00046919717,0.0003583178,0.0000067840274,0.00013405444,0.0020586334,0.000005325515,0.003837054],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9990395,0.00006970136,0.000067641195,0.00011074989,0.00039585677,0.00031657942],"domain_scores_gemma":[0.997758,0.00014105787,0.00027752313,0.00004479641,0.0012051869,0.00057342654],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0013795306,0.00035932532,0.0004170177,0.0022298512,0.005515558,0.0016148516,0.00093437126,0.00054106826,0.0030979365],"category_scores_gemma":[0.0026428367,0.00047092716,0.00090054923,0.005484258,0.0007638096,0.00052846805,0.0012193664,0.00085732865,0.000500546],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010170533,0.00013373785,0.93059,0.0001689881,0.000037937654,0.00018368849,0.049191102,0.000060894625,0.00039362616,0.0003411707,0.00274283,0.016054282],"study_design_scores_gemma":[0.0000035320359,0.000046137462,0.9688844,0.0000715238,0.00001014251,0.000027579143,0.028650519,0.00007158535,0.000035838704,0.000016211874,0.002166249,0.000016300102],"about_ca_topic_score_codex":0.99164575,"about_ca_topic_score_gemma":0.99654937,"teacher_disagreement_score":0.029346725,"about_ca_system_score_codex":0.029346725,"about_ca_system_score_gemma":0.035230182,"threshold_uncertainty_score":0.21292639},"labels":[],"label_agreement":null},{"id":"W7017453072","doi":"","title":"Automatic taxonomy evaluation","year":2022,"lang":"en","type":"dissertation","venue":"Papyrus : Institutional Repository (Université de Montréal)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Mitacs","keywords":"Taxonomy (biology); Demise","score_opus":0.01125728022391643,"score_gpt":0.21245689236071622,"score_spread":0.2011996121367998,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017453072","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.043369737,0.005073331,0.6748665,0.0014294273,0.0018295822,0.0029570232,0.052580908,0.09992411,0.11796937],"genre_scores_gemma":[0.12353701,0.001902297,0.720087,0.00063045847,0.0001785931,0.002045156,0.091258764,0.004498756,0.05586203],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99462587,0.0010739304,0.0005654145,0.0010437758,0.0022606272,0.00043028704],"domain_scores_gemma":[0.99177647,0.002010287,0.0003747775,0.00121576,0.0043559493,0.00026684196],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003597294,0.0025721914,0.0014166988,0.008198721,0.0020192272,0.005514771,0.0021955203,0.0019476209,0.053117044],"category_scores_gemma":[0.015696635,0.0007624872,0.0017251775,0.0059356787,0.0006427749,0.0065038875,0.0034555148,0.001532207,0.03227007],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040384865,0.0001319653,0.0044688326,0.0014007279,0.00011536504,0.00020278445,0.0006714376,0.0034517453,0.014121981,0.015716163,0.14564177,0.8136733],"study_design_scores_gemma":[0.00017914141,0.00037365127,0.017340869,0.00087753264,0.00019500381,0.0010290773,0.002995101,0.18022187,0.04592286,0.04344225,0.7072149,0.00020781498],"about_ca_topic_score_codex":0.011049566,"about_ca_topic_score_gemma":0.019977424,"teacher_disagreement_score":0.053117044,"about_ca_system_score_codex":0.0027461357,"about_ca_system_score_gemma":0.0043327753,"threshold_uncertainty_score":0.1776942},"labels":[],"label_agreement":null},{"id":"W7017481316","doi":"","title":"Bayesian Quantile Regression Based on the Generalized Gamma Distribution","year":2021,"lang":"en","type":"article","venue":"Scholarship at UWindsor (University of Windsor)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"University of Windsor","keywords":"Quantile regression; Akaike information criterion; Quantile; Deviance (statistics); Deviance information criterion; Bayesian probability; Bayesian information criterion; Bayesian linear regression; Gamma distribution; Bayes estimator","score_opus":0.021155143145881156,"score_gpt":0.24429700703777596,"score_spread":0.2231418638918948,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017481316","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005876402,0.00033117915,0.99254894,0.00015930009,0.000017373062,0.00003912456,0.00011247737,0.000270258,0.0006449341],"genre_scores_gemma":[0.5121278,0.0026429466,0.47653356,0.0003922622,0.0001647248,0.0005304006,0.0011687508,0.0004419065,0.0059975954],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99418193,0.0037749028,0.00017550518,0.00079355313,0.00077330833,0.0003008081],"domain_scores_gemma":[0.9872751,0.009828878,0.00092083,0.0008843594,0.00095582573,0.00013507195],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.012283274,0.0009909482,0.0018938503,0.0023618406,0.0005743793,0.0021161404,0.0023282631,0.0014236972,0.0040103924],"category_scores_gemma":[0.033606056,0.0006649167,0.0017205233,0.0031229937,0.0015989953,0.0022392343,0.0017801252,0.0024092314,0.0011040984],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002162316,0.00007664146,0.009966827,0.00026377983,0.00029437194,0.0002636893,0.00035806088,0.6309931,0.0024033834,0.18736304,0.0036203966,0.16418044],"study_design_scores_gemma":[0.000025338819,0.000045687888,0.0024757525,0.000060616785,0.00004435538,0.00011139738,0.000050610044,0.8968858,0.0006014117,0.09702228,0.0026331951,0.00004364977],"about_ca_topic_score_codex":0.008488133,"about_ca_topic_score_gemma":0.0052483175,"teacher_disagreement_score":0.012283274,"about_ca_system_score_codex":0.0015061508,"about_ca_system_score_gemma":0.0016786221,"threshold_uncertainty_score":0.06496096},"labels":[],"label_agreement":null},{"id":"W7017565989","doi":"","title":"canada 1-888-879-0163 norton antivirus tech support phone number","year":2016,"lang":"en","type":"other","venue":"OSF Preprints (OSF Preprints)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Phone; The Internet; Mobile phone; Service (business); Mobile phone tracking; Internet appliance","score_opus":0.009344229051988957,"score_gpt":0.25690388757232824,"score_spread":0.2475596585203393,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017565989","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0006438722,0.00024132004,0.0005239136,0.00040174866,0.00026514867,0.000048737966,0.0010432833,0.0014083836,0.9954236],"genre_scores_gemma":[0.0009809015,0.00012277579,0.00017276077,0.00008155848,0.000020369564,0.000009021371,0.00044345667,0.00020876701,0.99796045],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999551,0.00002553308,0.00001981119,0.00011637443,0.00022321388,0.00006414762],"domain_scores_gemma":[0.9987348,0.0001000586,0.000047280788,0.000107482214,0.0006169657,0.00039345294],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00028083278,0.0010404509,0.00061083824,0.0010752836,0.0017893644,0.004387066,0.0009281372,0.001445603,0.8896657],"category_scores_gemma":[0.0012841865,0.000556338,0.00033944804,0.0011557053,0.00031408202,0.0021020884,0.0014369156,0.0008976597,0.845548],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007148048,0.00007454474,0.0004900767,0.00009597867,0.0000030640035,0.00018524012,0.000084235435,0.00009291462,0.0017870786,0.0033991928,0.89361614,0.10010001],"study_design_scores_gemma":[0.000011081009,0.000025967476,0.00061571033,0.00005780035,0.0000040427003,0.00013763586,0.000119295335,0.00022226831,0.0003893716,0.00022996248,0.99817884,0.000008080256],"about_ca_topic_score_codex":0.021099288,"about_ca_topic_score_gemma":0.047957763,"teacher_disagreement_score":0.11033428,"about_ca_system_score_codex":0.0013261285,"about_ca_system_score_gemma":0.0017852419,"threshold_uncertainty_score":0.15737838},"labels":[],"label_agreement":null},{"id":"W7017976601","doi":"","title":"A Conferência de Montreal","year":2013,"lang":"pt","type":"other","venue":"Institutional Digital Library (Federal Senate)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Nucleofection; Gestational period; TSG101; Dysgeusia; Liquation; Diafiltration; Emperipolesis; Triacetin; Durvalumab","score_opus":0.008980927268787726,"score_gpt":0.21484452670307472,"score_spread":0.205863599434287,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7017976601","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0015990398,0.007468353,0.0026843774,0.0037264337,0.0033107449,0.00007170015,0.0061760223,0.0011507066,0.97381276],"genre_scores_gemma":[0.0039070505,0.0023035207,0.0011113043,0.00015388132,0.00019199702,0.00002678901,0.0012891529,0.0002747394,0.9907415],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99880123,0.00015430144,0.0000301194,0.0002702543,0.00055553677,0.00018846667],"domain_scores_gemma":[0.9989994,0.00013828326,0.000028525832,0.0001672661,0.00043877104,0.00022769067],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009769929,0.0013699334,0.00089712354,0.0024076677,0.003680715,0.0061786934,0.0013439142,0.0012702455,0.5288616],"category_scores_gemma":[0.0020720765,0.0004119802,0.0007938055,0.003762828,0.0013277315,0.0016684717,0.0024364104,0.0013218996,0.15799373],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000071226736,0.00005457666,0.00043806343,0.0001694809,0.000013048365,0.00017226774,0.00025931204,0.000319312,0.00075900526,0.048080187,0.78204316,0.1676203],"study_design_scores_gemma":[0.0000048021097,0.0000047160415,0.00052675087,0.000036623765,0.0000034174252,0.000025984556,0.00006482367,0.00007734877,0.00013771889,0.0011829955,0.99792993,0.0000047826766],"about_ca_topic_score_codex":0.28900883,"about_ca_topic_score_gemma":0.4730288,"teacher_disagreement_score":0.5288616,"about_ca_system_score_codex":0.006737251,"about_ca_system_score_gemma":0.011764642,"threshold_uncertainty_score":0.67202175},"labels":[],"label_agreement":null},{"id":"W7019024081","doi":"","title":"Exploring the Role of Terminology in SNOMED CT Definitions: Challenges and Solutions","year":2023,"lang":"en","type":"article","venue":"Lirias (KU Leuven)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"Fundação para a Ciência e a Tecnologia; Centre hospitalier universitaire Sainte-Justine","keywords":"SNOMED CT; Terminology; Representation (politics); Nucleofection; Subject (documents); Knowledge representation and reasoning","score_opus":0.22378799338582336,"score_gpt":0.2839765758173448,"score_spread":0.06018858243152145,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7019024081","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013815544,0.09457531,0.49478,0.37174168,0.0051513235,0.00056304166,0.00087978237,0.0007681876,0.017725198],"genre_scores_gemma":[0.1405077,0.06997259,0.7309965,0.042863727,0.0060239118,0.0011563571,0.0024772496,0.0010115205,0.004990399],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.8780721,0.07783283,0.009726617,0.010053608,0.02206961,0.0022451784],"domain_scores_gemma":[0.71773976,0.2126912,0.014425445,0.01701546,0.034023676,0.0041044634],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.113735326,0.0024917747,0.00368357,0.011114516,0.0066389455,0.027298236,0.011751171,0.010838455,0.0055767223],"category_scores_gemma":[0.15726241,0.0022248174,0.0023565644,0.014870521,0.022712603,0.06735675,0.022642894,0.017009769,0.0025167838],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014033563,0.00021398852,0.0045399833,0.0069676666,0.00014131454,0.0014467282,0.018479861,0.0040919567,0.0025237093,0.67405164,0.037454408,0.24994835],"study_design_scores_gemma":[0.000025060874,0.00008452425,0.0008313505,0.0083429245,0.00009414678,0.0021097283,0.024385422,0.015822312,0.0015433345,0.7195722,0.22696471,0.00022425047],"about_ca_topic_score_codex":0.008772492,"about_ca_topic_score_gemma":0.008911021,"teacher_disagreement_score":0.113735326,"about_ca_system_score_codex":0.012560678,"about_ca_system_score_gemma":0.022241931,"threshold_uncertainty_score":0.60149705},"labels":[],"label_agreement":null},{"id":"W7024101912","doi":"","title":"Prepublication data sharing.","year":2009,"lang":"en","type":"article","venue":"Oxford University Research Archive (ORA) (University of Oxford)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Field (mathematics); Dozen; Data collection; Natural (archaeology); Biological organism; Context (archaeology); Timeline","score_opus":0.06052001179200812,"score_gpt":0.30584692541935526,"score_spread":0.24532691362734715,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024101912","genre_codex":"dataset","genre_gemma":"commentary","domain_codex":null,"domain_gemma":"reproducibility","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":null,"domain_candidate":"reproducibility","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0041905325,0.0012870382,0.09971035,0.004109851,0.005539821,0.0022788234,0.7815672,0.056162927,0.045153424],"genre_scores_gemma":[0.012834176,0.0010733597,0.09585226,0.0021199237,0.0005966173,0.003106025,0.8402861,0.019649768,0.024481693],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97687465,0.005968616,0.006814518,0.0036318297,0.0051826024,0.0015277157],"domain_scores_gemma":[0.84010005,0.026706677,0.007952215,0.096655056,0.023220813,0.0053651864],"candidate_categories":["metaresearch","open_science"],"consensus_categories":[],"category_scores_codex":[0.029230386,0.0017001469,0.002631729,0.011651659,0.0030844356,0.0090343,0.0047125667,0.002992917,0.19429334],"category_scores_gemma":[0.096158154,0.0020762922,0.0025562064,0.02041941,0.0020278022,0.008331766,0.011148571,0.0046662167,0.18664575],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0013413581,0.00016356786,0.0024060362,0.00272185,0.00015920219,0.0004517423,0.0011371919,0.00045692804,0.008349459,0.010764792,0.85552496,0.116522856],"study_design_scores_gemma":[0.00012219718,0.000066691944,0.0019610934,0.00049267564,0.00006437834,0.0003330018,0.00033314727,0.00065941905,0.00824099,0.015774902,0.9718441,0.00010742599],"about_ca_topic_score_codex":0.0025224641,"about_ca_topic_score_gemma":0.0028795267,"teacher_disagreement_score":0.9952874,"about_ca_system_score_codex":0.0019431685,"about_ca_system_score_gemma":0.009788261,"threshold_uncertainty_score":0.64997596},"labels":[],"label_agreement":null},{"id":"W7024986682","doi":"","title":"This Week...Public Sector Employment; Farm Income; Remand in Correctional Facilities; Unemployment and Employment Insurance; Consumer Price Inflation","year":2011,"lang":"en","type":"report","venue":"oURspace (University of Regina)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Unemployment; Government (linguistics); Public policy; Inflation (cosmology); Point (geometry); Monetary policy; Empirical evidence","score_opus":0.03573616753177251,"score_gpt":0.24338197296911077,"score_spread":0.20764580543733827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7024986682","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012024243,0.0017842571,0.00077896327,0.05896461,0.009245307,0.00039530243,0.1989095,0.0023595695,0.71553826],"genre_scores_gemma":[0.009414455,0.0011213849,0.00043722717,0.001389338,0.0005799322,0.00006946472,0.027729392,0.00019297007,0.9590659],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994717,0.00002017549,0.000021937914,0.000030869516,0.00032311154,0.00013215173],"domain_scores_gemma":[0.99793386,0.0001427681,0.00010420364,0.00008761798,0.0009944551,0.00073708413],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00062490936,0.0005268232,0.0002466111,0.0018491554,0.0020161709,0.003548735,0.00064757414,0.0009344255,0.21480751],"category_scores_gemma":[0.001964123,0.0002491097,0.0002844081,0.002818108,0.0003755773,0.0009551625,0.0010693969,0.0010016508,0.10065068],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007845967,0.000018095277,0.0019879558,0.00002026747,8.479626e-7,0.000023046905,0.00008111706,0.000012031061,0.00003347563,0.0002784911,0.9807569,0.01677993],"study_design_scores_gemma":[0.0000048999373,0.000011216315,0.034702245,0.000069967864,0.000002650684,0.00004257158,0.00060714077,0.000066743996,0.00015962526,0.00017675021,0.96414787,0.000008394586],"about_ca_topic_score_codex":0.4716416,"about_ca_topic_score_gemma":0.73120964,"teacher_disagreement_score":0.5283584,"about_ca_system_score_codex":0.0036324833,"about_ca_system_score_gemma":0.006875766,"threshold_uncertainty_score":0.93779266},"labels":[],"label_agreement":null},{"id":"W7027337249","doi":"","title":"CAETS Forum, Session 3: noise transmission in buildings","year":2009,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Noise (video); Transmission (telecommunications); Session (web analytics); Background noise; Mode (computer interface)","score_opus":0.00934992368487664,"score_gpt":0.2684319495571851,"score_spread":0.25908202587230844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7027337249","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.028946664,0.015745625,0.061509483,0.14792295,0.2628039,0.0044714636,0.05724262,0.00949709,0.4118603],"genre_scores_gemma":[0.0626252,0.004257584,0.013868811,0.00551162,0.044344638,0.0014790294,0.042541932,0.0024666798,0.8229045],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9974388,0.0005898331,0.00014048725,0.00030550334,0.001224723,0.00030056995],"domain_scores_gemma":[0.98827016,0.0024922467,0.00044018362,0.0009851642,0.0053726747,0.0024396016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0068510026,0.0011412504,0.0012137671,0.0021054999,0.0035908103,0.0062005604,0.002350977,0.0061879824,0.20488048],"category_scores_gemma":[0.010810512,0.0003665319,0.0010791498,0.0016417587,0.0006512825,0.0045074727,0.0039212084,0.0030439212,0.051107984],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002680578,0.00008578982,0.0002765584,0.0002046344,0.00001323583,0.00010688798,0.00008413619,0.00037769414,0.0011670882,0.0015082264,0.9734807,0.022426985],"study_design_scores_gemma":[0.000083396415,0.00013659589,0.0017306114,0.00016052606,0.00002278138,0.0000805706,0.00046068506,0.0016549382,0.0018033178,0.0030718034,0.9907624,0.00003227236],"about_ca_topic_score_codex":0.0039792885,"about_ca_topic_score_gemma":0.008550129,"teacher_disagreement_score":0.20488048,"about_ca_system_score_codex":0.001174202,"about_ca_system_score_gemma":0.0031849018,"threshold_uncertainty_score":0.68539345},"labels":[],"label_agreement":null},{"id":"W7028320348","doi":"","title":"Falling Gasoline Use Means United States Can Just Say No to New Pipelines and Food-to-Fuel","year":2013,"lang":"en","type":"dataset","venue":"Issue Lab (Candid)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Pipeline transport; Oil refinery; Pipeline (software); Falling (accident); Oil sands; Fossil fuel; Petroleum","score_opus":0.029051572687169923,"score_gpt":0.28342855780070875,"score_spread":0.25437698511353884,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7028320348","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0019277952,0.00016110757,0.000183643,0.00043078867,0.000050155053,0.000013771729,0.99479926,0.0005078025,0.00192565],"genre_scores_gemma":[0.0021557303,0.00011773648,0.0005016565,0.000102726655,0.000008412133,0.000026541904,0.99613404,0.000029815948,0.0009232844],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.999185,0.000089782545,0.000100646466,0.0002180814,0.00028251988,0.00012395864],"domain_scores_gemma":[0.9985536,0.00034704412,0.00019852602,0.00032051708,0.00037457648,0.00020573453],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000827305,0.0014919157,0.00073506276,0.002719049,0.0009746216,0.0016374267,0.0021289466,0.0016288947,0.011822399],"category_scores_gemma":[0.004031949,0.00035663813,0.001004972,0.0066064396,0.0005591361,0.0012642973,0.0015256018,0.0018022576,0.014181082],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00014202057,0.00006229349,0.0042444575,0.00042118688,0.000039609942,0.000059933416,0.000039484123,0.0004978371,0.00027222358,0.000890925,0.98698425,0.006345698],"study_design_scores_gemma":[0.00021478692,0.000032044678,0.01843984,0.00020863874,0.00004978299,0.00021668467,0.00022049382,0.0020117771,0.0015436949,0.0014416516,0.97558165,0.000039046354],"about_ca_topic_score_codex":0.077596344,"about_ca_topic_score_gemma":0.17271374,"teacher_disagreement_score":0.077596344,"about_ca_system_score_codex":0.0024381483,"about_ca_system_score_gemma":0.0027027798,"threshold_uncertainty_score":0.15428936},"labels":[],"label_agreement":null},{"id":"W7028483255","doi":"","title":"Free and Always Will Be? On Social Media Participation as it Undermines Individual Autonomy","year":2021,"lang":"en","type":"article","venue":"PhilPapers (PhilPapers Foundation)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Trent University","funders":"","keywords":"Social media; Autonomy; Harassment; Formative assessment; Misinformation; Social relation; Digital media","score_opus":0.05350515600983685,"score_gpt":0.31083374331432917,"score_spread":0.2573285873044923,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7028483255","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.21983182,0.0025494883,0.016131643,0.31468233,0.0012469557,0.00007095818,0.00007671377,0.00016325331,0.4452468],"genre_scores_gemma":[0.9636741,0.0008068158,0.001541971,0.018836139,0.0003775279,0.00007368359,0.000023194923,0.00008717084,0.014579407],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9734178,0.01747941,0.00047661952,0.0020295419,0.0042557865,0.002340935],"domain_scores_gemma":[0.9729442,0.017064434,0.00213483,0.0038510067,0.0015384743,0.0024670383],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020566,0.00040028238,0.00056182616,0.001326992,0.010955381,0.018167002,0.0020224864,0.007254349,0.00647],"category_scores_gemma":[0.030651134,0.00033946763,0.0006584677,0.0009178507,0.06460326,0.026180955,0.01759157,0.008038034,0.0011078786],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000052821364,0.000049531056,0.0023952445,0.00005240346,0.000015746997,0.0005728226,0.15762666,0.00012369706,0.00070484186,0.79475975,0.0115716,0.03207489],"study_design_scores_gemma":[0.000030013263,0.00009921352,0.0026819282,0.000538532,0.000026647063,0.0007119291,0.13587147,0.00057117234,0.0010660437,0.5787659,0.2795599,0.00007722992],"about_ca_topic_score_codex":0.0026719207,"about_ca_topic_score_gemma":0.003017677,"teacher_disagreement_score":0.020566,"about_ca_system_score_codex":0.003581645,"about_ca_system_score_gemma":0.0045491834,"threshold_uncertainty_score":0.10876471},"labels":[],"label_agreement":null},{"id":"W7028882081","doi":"","title":"Genetic mapping of quantitative trait loci influencing growth, development and morphology in Atlantic salmon (Salmo salar, L.)","year":2003,"lang":"en","type":"other","venue":"OpenGrey (Institut de l'Information Scientifique et Technique)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Quantitative trait locus; Salmo; Backcrossing; Morphology (biology); Phenotype; Genetic variation; Quantitative genetics; Phenotypic trait; Trait; Adaptation (eye)","score_opus":0.02142778047727205,"score_gpt":0.2722821648471088,"score_spread":0.25085438436983676,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7028882081","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992085,0.00007423079,0.00033780996,0.0000140020065,0.000001556231,0.000008227745,0.0001489671,0.000006266768,0.00020049003],"genre_scores_gemma":[0.9967989,0.00016504718,0.0012858013,0.000039643714,0.0000031736038,0.000038439397,0.00054243015,0.000010705206,0.0011157432],"study_design_codex":"bench_or_experimental","study_design_gemma":"observational","domain_scores_codex":[0.99987304,0.000020030591,0.0000070702104,0.000039935854,0.000035521174,0.000024403244],"domain_scores_gemma":[0.99967825,0.000117327836,0.00011147281,0.00001126584,0.00003517147,0.000046540892],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00026483307,0.00022801048,0.00016693046,0.00056655536,0.00017499078,0.0002118696,0.0001970448,0.00018991559,0.00037742773],"category_scores_gemma":[0.0002222226,0.00022300902,0.00027233112,0.00032535393,0.00028044963,0.000069146954,0.0002759092,0.00025192238,0.00013677226],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00024921334,0.00005692417,0.01531485,0.00003210559,0.000021237078,0.000085411426,0.00019070892,0.00018169514,0.9797309,0.00006546061,0.00002045301,0.004050934],"study_design_scores_gemma":[0.000052705185,0.00075651176,0.9684267,0.000012027834,0.00005006395,0.00017603213,0.00021011624,0.00076860836,0.0290164,0.00005719921,0.00045874107,0.0000148152],"about_ca_topic_score_codex":0.012006332,"about_ca_topic_score_gemma":0.029169358,"teacher_disagreement_score":0.012006332,"about_ca_system_score_codex":0.0004683551,"about_ca_system_score_gemma":0.0005063834,"threshold_uncertainty_score":0.023872912},"labels":[],"label_agreement":null},{"id":"W7033045855","doi":"","title":"A press full of pamphlets on Ireland, stereotypes, sensationalism, and veracity in English reactions to the 1641 Irish rebellion, November 1641-August 1642","year":2001,"lang":"en","type":"other","venue":"Library and Archives Canada (Government of Canada)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Irish; Period (music); Government (linguistics); Subject (documents)","score_opus":0.0060164247029340905,"score_gpt":0.17901437260061312,"score_spread":0.17299794789767903,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7033045855","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0027446726,0.0064250277,0.00037459223,0.090898745,0.016356999,0.000110571054,0.004270471,0.0005253986,0.87829363],"genre_scores_gemma":[0.004159553,0.0014191786,0.00015124305,0.0065503637,0.0010622629,0.00003603777,0.00066797645,0.00019825634,0.9857552],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994367,0.00009045842,0.00002871202,0.00004304738,0.0002660042,0.00013506136],"domain_scores_gemma":[0.9965063,0.0016001181,0.0001473412,0.00015419372,0.00080981344,0.00078225875],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0015743086,0.0005263062,0.00046099463,0.0024504354,0.006040392,0.0059831827,0.00074338616,0.003096572,0.20050828],"category_scores_gemma":[0.005265016,0.00040069912,0.00020604592,0.004865549,0.0024451134,0.004271776,0.0022087912,0.0040751896,0.04533349],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000007541191,0.0000058463174,0.000065469874,0.000028811788,4.187403e-7,0.000037704158,0.0005846432,0.000008091165,0.000026885726,0.0013061088,0.9925827,0.005345884],"study_design_scores_gemma":[0.0000045194943,0.0000042939823,0.0020490445,0.00009931251,0.0000013053261,0.000052489613,0.003268511,0.000014973441,0.00003829638,0.00040509822,0.99405485,0.00000742194],"about_ca_topic_score_codex":0.0985716,"about_ca_topic_score_gemma":0.33977753,"teacher_disagreement_score":0.20050828,"about_ca_system_score_codex":0.004691171,"about_ca_system_score_gemma":0.0058609946,"threshold_uncertainty_score":0.67076707},"labels":[],"label_agreement":null},{"id":"W7033772052","doi":"","title":"THE ROLE OF CODE SWITCHING PHENOMENA IN A YOUTUBE VLOG&#13;\\nBY SACHA STEVENSON&#13;\\n","year":2019,"lang":"en","type":"dissertation","venue":"UNDIP Institutional Repository (UNDIP-IR) (Diponegoro University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Code-switching; Upload; Code (set theory); Channel (broadcasting); Function (biology); Social media","score_opus":0.006714281022849123,"score_gpt":0.21155566865765355,"score_spread":0.20484138763480442,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7033772052","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9020234,0.00051041675,0.0041540163,0.003687708,0.00010433121,0.00015430598,0.00019744903,0.000114596565,0.08905389],"genre_scores_gemma":[0.9955901,0.00016911475,0.0008735232,0.0002500317,0.00001772538,0.000048361424,0.000106271975,0.000048509704,0.002896343],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9958974,0.0021265429,0.0001289191,0.00045045558,0.0010326811,0.00036407218],"domain_scores_gemma":[0.98722315,0.007494227,0.00211872,0.0006887995,0.00157147,0.0009036901],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0030637619,0.00020135252,0.0001557817,0.002389036,0.0045141834,0.004355377,0.0008672948,0.0008397483,0.0028691879],"category_scores_gemma":[0.018838575,0.00022135876,0.00023111013,0.0018381377,0.0077172737,0.0055351355,0.0033752795,0.0016027411,0.00028362122],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00019662004,0.00012238327,0.07617374,0.00020541369,0.00002747364,0.0021389676,0.7677682,0.00023409718,0.0069988496,0.07088916,0.008976015,0.06626899],"study_design_scores_gemma":[0.000028986111,0.00017813897,0.19354677,0.00054773106,0.000028881144,0.0010317157,0.7062275,0.0038963086,0.0038390981,0.015635112,0.07491108,0.00012872835],"about_ca_topic_score_codex":0.04088925,"about_ca_topic_score_gemma":0.034260903,"teacher_disagreement_score":0.04088925,"about_ca_system_score_codex":0.0043053566,"about_ca_system_score_gemma":0.0019880647,"threshold_uncertainty_score":0.08130252},"labels":[],"label_agreement":null},{"id":"W7033924084","doi":"","title":"S1E22 - Goldie Morgentaler: Keeping Yiddishkeit alive in Lethbridge","year":2022,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Derogation; Natural (archaeology); Pretext; Work (physics)","score_opus":0.007086301888271251,"score_gpt":0.2021649795842188,"score_spread":0.19507867769594756,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7033924084","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.000993289,0.0027705135,0.0034697042,0.04515566,0.018178724,0.000343062,0.05874932,0.01745926,0.8528805],"genre_scores_gemma":[0.0010739465,0.00050637685,0.0009654758,0.0028265768,0.00069846614,0.00008514133,0.011211882,0.0020965738,0.9805355],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990565,0.00013704899,0.000036002297,0.000112481284,0.0005042653,0.00015368518],"domain_scores_gemma":[0.9957736,0.00066890917,0.00009840025,0.00025975547,0.0017400989,0.0014591841],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0019034666,0.00083727465,0.0006989714,0.0023790717,0.0018004228,0.0034934413,0.0013982301,0.002144953,0.8116562],"category_scores_gemma":[0.008679374,0.00041788226,0.00036873252,0.002904482,0.0005694366,0.0025347327,0.00401012,0.0017235818,0.6253394],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000059591316,0.0000026752364,0.000021520917,0.000013139925,2.4317902e-7,0.000012331839,0.000016173362,0.000004261229,0.000028149067,0.00011766209,0.99069715,0.009080761],"study_design_scores_gemma":[0.0000033170375,0.0000027184585,0.00018274767,0.000025427069,5.9524774e-7,0.000013548796,0.000043504217,0.000019874667,0.00004182523,0.00016172475,0.9995018,0.0000029051266],"about_ca_topic_score_codex":0.02310889,"about_ca_topic_score_gemma":0.067658484,"teacher_disagreement_score":0.18834382,"about_ca_system_score_codex":0.0016825796,"about_ca_system_score_gemma":0.0037089912,"threshold_uncertainty_score":0.26864964},"labels":[],"label_agreement":null},{"id":"W7034324606","doi":"","title":"Space Exploration Workbook (Grade 9 Ontario Science SNC1W)","year":2023,"lang":"en","type":"other","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Workbook; Space (punctuation); Data exploration; Space exploration; Space research","score_opus":0.03512000505504143,"score_gpt":0.28761760000158576,"score_spread":0.25249759494654433,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7034324606","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005624296,0.00040222507,0.009516446,0.0011239675,0.000648023,0.00017079899,0.045242395,0.014175056,0.92815864],"genre_scores_gemma":[0.0023752989,0.00049836194,0.008143768,0.000256604,0.00015043709,0.00015939104,0.03134608,0.007831501,0.94923854],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992587,0.000054583106,0.000022825194,0.00007242392,0.00050698215,0.00008439119],"domain_scores_gemma":[0.9972373,0.00031414028,0.000061791085,0.00033594898,0.001512068,0.00053869473],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0009191052,0.0014085739,0.0008671031,0.0044857925,0.0018463009,0.004590557,0.002045871,0.0010455158,0.82457536],"category_scores_gemma":[0.0040515237,0.000693205,0.00086518755,0.007333702,0.00062427536,0.0030879504,0.003740762,0.0010614571,0.62621284],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018704453,0.000008514019,0.00006085842,0.000057900575,0.000001149271,0.000011316415,0.000035436467,0.00016324421,0.00013641217,0.0019201073,0.95177585,0.045810565],"study_design_scores_gemma":[0.0000068335053,0.0000034107804,0.00015183983,0.000025480613,0.0000012151935,0.000008842605,0.00003536,0.00022663415,0.00013048293,0.0018582473,0.9975466,0.0000050197086],"about_ca_topic_score_codex":0.08369199,"about_ca_topic_score_gemma":0.22267157,"teacher_disagreement_score":0.916308,"about_ca_system_score_codex":0.0034737082,"about_ca_system_score_gemma":0.0065201386,"threshold_uncertainty_score":0.25022197},"labels":[],"label_agreement":null},{"id":"W7034830233","doi":"","title":"Watch Boys State (2020) Full Movie Online Free","year":2020,"lang":"en","type":"other","venue":"OSF Preprints (OSF Preprints)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Drama; State (computer science); Action (physics); Plot (graphics); Identity (music); Shot (pellet); Worry; Stalking; Call to action","score_opus":0.01216204691221999,"score_gpt":0.25857433303846283,"score_spread":0.24641228612624283,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7034830233","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00089013175,0.00028745725,0.0005196682,0.0008191719,0.0012070878,0.00015088287,0.0072635617,0.005592216,0.98326993],"genre_scores_gemma":[0.004778136,0.00042723928,0.0006001092,0.00071504264,0.00052140857,0.00011877449,0.005212341,0.0025985783,0.9850285],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9998367,0.00001473443,0.0000049720047,0.00002838132,0.00006265039,0.000052521325],"domain_scores_gemma":[0.9992749,0.000054871052,0.000021163716,0.000056425273,0.00020110619,0.0003914737],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00024454342,0.0007951629,0.0004342431,0.00080973125,0.0014235257,0.0027824105,0.0006508122,0.0011068017,0.88995737],"category_scores_gemma":[0.001041693,0.0004237896,0.0005176584,0.0005647396,0.0002768515,0.0030890282,0.0029382315,0.0012507555,0.75674516],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000028202301,0.000012223044,0.00008542126,0.00004734641,0.0000017015269,0.000026074085,0.000049965824,0.000009450534,0.00017190362,0.00035376666,0.9871192,0.012094712],"study_design_scores_gemma":[0.000015536256,0.000016889748,0.0007992763,0.00004814388,0.0000026906664,0.000035678888,0.00010413538,0.000043233547,0.00014616447,0.00013412142,0.9986469,0.0000071485965],"about_ca_topic_score_codex":0.0036536304,"about_ca_topic_score_gemma":0.016205532,"teacher_disagreement_score":0.11004263,"about_ca_system_score_codex":0.00049042265,"about_ca_system_score_gemma":0.00035190847,"threshold_uncertainty_score":0.15696245},"labels":[],"label_agreement":null},{"id":"W7035707450","doi":"","title":"Accepting alternate solutions under objective-based codes","year":2001,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Term (time); Power (physics); Order (exchange); Work (physics)","score_opus":0.036285697293626996,"score_gpt":0.2960232257409393,"score_spread":0.2597375284473123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7035707450","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.08966599,0.00025105884,0.87014043,0.0023518885,0.00022001904,0.00048322807,0.00039959172,0.0013836955,0.035104107],"genre_scores_gemma":[0.3460514,0.00025213623,0.6361178,0.00050930073,0.0001256074,0.0009013442,0.0008465585,0.00052511104,0.01467066],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9886035,0.0033045607,0.0009611185,0.0013439703,0.0047426466,0.0010442304],"domain_scores_gemma":[0.97135204,0.013723546,0.0021154583,0.005626119,0.0063043716,0.00087836466],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008986919,0.0010309569,0.0009616136,0.0020501483,0.0015708162,0.005170372,0.0025353718,0.002465092,0.0103823105],"category_scores_gemma":[0.04717112,0.00046016838,0.0014025965,0.001678181,0.0025912789,0.006859427,0.0054686805,0.0024698984,0.0018724629],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006567692,0.00034876773,0.0060111023,0.0005597991,0.00013481524,0.0004205528,0.0024810606,0.09951608,0.0050126687,0.5711313,0.005870195,0.3078569],"study_design_scores_gemma":[0.00016681573,0.00036586783,0.0010601386,0.00024157584,0.00008947851,0.00023498668,0.001325148,0.5010739,0.0075514275,0.4634407,0.024370153,0.00007989432],"about_ca_topic_score_codex":0.0031636592,"about_ca_topic_score_gemma":0.005472365,"teacher_disagreement_score":0.0103823105,"about_ca_system_score_codex":0.0025652351,"about_ca_system_score_gemma":0.0036546686,"threshold_uncertainty_score":0.04752797},"labels":[],"label_agreement":null},{"id":"W7035743767","doi":"","title":"Angel:::1844-234-9752 Netgear router tech support phone number(1844-234-9752) netgear router technical support phone number 1844::234::9752","year":2016,"lang":"en","type":"other","venue":"OSF Preprints (OSF Preprints)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Core router; Router; One-armed router; Phone; Wireless network","score_opus":0.011688443262062006,"score_gpt":0.27799354619185074,"score_spread":0.2663051029297887,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7035743767","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005656113,0.00011139337,0.0031617973,0.00043584866,0.0008319369,0.00012172082,0.0019355163,0.010373082,0.9824631],"genre_scores_gemma":[0.0020782398,0.00016866709,0.0009659258,0.00038236877,0.00013297227,0.00006847072,0.0016350123,0.002051871,0.99251646],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990927,0.00007068968,0.000031026746,0.00024981017,0.0003600048,0.00019572122],"domain_scores_gemma":[0.99824274,0.0001275039,0.000053269207,0.00020404873,0.00073186134,0.00064056576],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00064266653,0.0014891968,0.0012245015,0.00117177,0.001826607,0.006488014,0.0014049507,0.0019341119,0.90878665],"category_scores_gemma":[0.0019267584,0.000811717,0.00068296544,0.0010670547,0.00032587422,0.0029777763,0.0024185253,0.0022624957,0.91559786],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006752815,0.00005337206,0.0001920558,0.00004562038,0.000005118475,0.00004060611,0.00003950755,0.00017250479,0.0013196327,0.002959905,0.9534675,0.041636627],"study_design_scores_gemma":[0.000030047644,0.00004169829,0.0004180787,0.00003918025,0.000009486595,0.00006359764,0.00006873326,0.0003939214,0.0008579388,0.0005378196,0.9975247,0.000014819554],"about_ca_topic_score_codex":0.004354455,"about_ca_topic_score_gemma":0.005017399,"teacher_disagreement_score":0.091213346,"about_ca_system_score_codex":0.0016547171,"about_ca_system_score_gemma":0.0016589242,"threshold_uncertainty_score":0.13010466},"labels":[],"label_agreement":null},{"id":"W7035987659","doi":"","title":"Active Shooter 101 (26 HQ PDFs &amp; 20 HD Videos)","year":2015,"lang":"en","type":"other","venue":"Bulletin of Miscellaneous Information (Royal Gardens Kew)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Service (business); Active monitoring; Summit; Event (particle physics); Poison control; Fire investigation","score_opus":0.010102015402189666,"score_gpt":0.21724847937952194,"score_spread":0.20714646397733227,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7035987659","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00085818453,0.00030818366,0.0047362004,0.0009438275,0.0017249866,0.0007961171,0.05003898,0.030347418,0.910246],"genre_scores_gemma":[0.0031877575,0.00033793924,0.0029972182,0.0007886642,0.00058352837,0.00047816406,0.028468031,0.008612237,0.9545465],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996153,0.000046674682,0.000020380432,0.00006069706,0.00018815171,0.00006874567],"domain_scores_gemma":[0.99748945,0.00047241006,0.00009811063,0.00036292785,0.0010347059,0.00054240855],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00059247104,0.0012557284,0.0008230856,0.0019163276,0.0011557708,0.0035724402,0.0017190702,0.0012823885,0.9450599],"category_scores_gemma":[0.0040403237,0.0007841797,0.00058065844,0.0021232804,0.00033180616,0.003714747,0.002909586,0.0016540203,0.8670402],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003270787,0.000022461767,0.00002695227,0.00007323089,0.0000010142029,0.000012681679,0.00001978065,0.000023084827,0.00012703316,0.0002252506,0.9746626,0.024773166],"study_design_scores_gemma":[0.00004539502,0.000051692117,0.00071621116,0.000091009235,0.0000026896441,0.00007237498,0.00010452416,0.00011991162,0.00030726165,0.0005137023,0.99796283,0.000012354064],"about_ca_topic_score_codex":0.002948615,"about_ca_topic_score_gemma":0.004678679,"teacher_disagreement_score":0.054940104,"about_ca_system_score_codex":0.00063115085,"about_ca_system_score_gemma":0.00085072813,"threshold_uncertainty_score":0.078365326},"labels":[],"label_agreement":null},{"id":"W7036417148","doi":"","title":"Center on Disability Studies eNewsletter, March 2024","year":2024,"lang":"en","type":"article","venue":"ScholarSpace (University of Hawaii at Manoa)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Center (category theory); Join (topology); Disability studies; Research center; Canadian studies","score_opus":0.027281781877012398,"score_gpt":0.2884622127270135,"score_spread":0.2611804308500011,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7036417148","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0010983107,0.015595482,0.00028458805,0.13035533,0.044709418,0.00023461359,0.027469158,0.0007849988,0.7794681],"genre_scores_gemma":[0.002843925,0.004164558,0.00019410907,0.007967957,0.002465417,0.00021118177,0.00528917,0.00014207906,0.97672164],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991505,0.0001081277,0.0000634308,0.00011795682,0.00036173215,0.0001982055],"domain_scores_gemma":[0.99709964,0.00029199923,0.000099613186,0.00016172456,0.0010211115,0.0013258422],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0015315339,0.000765475,0.00064826076,0.0014961233,0.0037060133,0.005106895,0.0010383283,0.0032747267,0.47129294],"category_scores_gemma":[0.0062487843,0.00043690167,0.0003153037,0.0018909917,0.00043137895,0.0029315576,0.004498929,0.0031445408,0.22074053],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000044039443,0.000005968899,0.000049802224,0.000009729627,2.3781126e-7,0.000011228056,0.000012552947,0.0000014125958,0.0000047774697,0.00027324172,0.9941273,0.0054994766],"study_design_scores_gemma":[0.0000042986467,0.0000037342552,0.0006759632,0.000066716064,5.115249e-7,0.0000108626455,0.00011616119,0.0000032963612,0.0000078085195,0.00012223961,0.99898666,0.0000018459232],"about_ca_topic_score_codex":0.02120925,"about_ca_topic_score_gemma":0.06499324,"teacher_disagreement_score":0.47129294,"about_ca_system_score_codex":0.0020463236,"about_ca_system_score_gemma":0.0060201734,"threshold_uncertainty_score":0.75413644},"labels":[],"label_agreement":null},{"id":"W7036614036","doi":"","title":"Chronic obstructive pulmonary disease (COPD): the Impact of occupational hazards in the minerals industry","year":2022,"lang":"en","type":"dissertation","venue":"Lu Zone Ul (Laurentian University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Compensation (psychology); Workers' compensation; COPD; Occupational disease; Pulmonary disease; Thematic analysis; Disease; Occupational medicine","score_opus":0.013199364892702606,"score_gpt":0.2723426604440922,"score_spread":0.2591432955513896,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7036614036","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98483974,0.0017048973,0.000078832476,0.0042412723,0.000038251634,0.000017077782,0.00008691626,0.000002590807,0.008990366],"genre_scores_gemma":[0.99699235,0.0011616319,0.00005223305,0.00039626515,0.000026135995,0.000009353667,0.000031451225,0.0000012288863,0.0013293462],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9990004,0.00029082698,0.000041032014,0.00007022666,0.00033394937,0.00026361924],"domain_scores_gemma":[0.999119,0.00023600487,0.00019503137,0.000016446602,0.00013756855,0.00029602827],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006496995,0.00010039968,0.00016942805,0.00041772772,0.0041100797,0.0018130193,0.00036178797,0.0005626338,0.002084701],"category_scores_gemma":[0.0014150713,0.00010477271,0.00015584796,0.00066422793,0.0019882445,0.0006150294,0.0026417999,0.00059214444,0.00011052303],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000096433236,0.00018730153,0.55734056,0.000636186,0.000032647047,0.0028935713,0.3923445,0.0001486037,0.0028765209,0.0015676435,0.004677525,0.037198484],"study_design_scores_gemma":[0.0000072598546,0.00008558022,0.65263104,0.0002905028,0.000015847396,0.00073094334,0.31959316,0.0000531732,0.00012363611,0.00020697116,0.026244383,0.000017537503],"about_ca_topic_score_codex":0.39744452,"about_ca_topic_score_gemma":0.6987764,"teacher_disagreement_score":0.39744452,"about_ca_system_score_codex":0.0061074332,"about_ca_system_score_gemma":0.009930093,"threshold_uncertainty_score":0.7902623},"labels":[],"label_agreement":null},{"id":"W7037379210","doi":"","title":"Effect of Economy and FDA Intervention on the Hearing Aid Industry","year":2005,"lang":"en","type":"article","venue":"Lancaster EPrints (Lancaster University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Quarter (Canadian coin); Hearing aid; Intervention (counseling); Ordinary least squares; Supply and demand; Hearing loss; Flourishing","score_opus":0.012900112967037348,"score_gpt":0.23604079066169545,"score_spread":0.2231406776946581,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037379210","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98417896,0.00087774755,0.00048800203,0.002711598,0.00006744007,0.00010229234,0.00033097455,0.00002808011,0.011214791],"genre_scores_gemma":[0.99690634,0.00022794439,0.00017453068,0.00064621144,0.00005394394,0.000057769288,0.00013062061,0.0000055166042,0.0017970959],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99653924,0.0018145096,0.00014028599,0.00024955126,0.0005516146,0.00070484955],"domain_scores_gemma":[0.9604019,0.025971498,0.009681822,0.00063918123,0.0016820336,0.0016235553],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0034026885,0.00027518123,0.00039437425,0.00061629835,0.00043908448,0.0012124452,0.00035550416,0.0008881404,0.012584949],"category_scores_gemma":[0.019481968,0.00013468291,0.00043954534,0.00040622285,0.00088137016,0.00048788407,0.0011159818,0.0011966291,0.000761827],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.018131804,0.007037812,0.7953977,0.0005074648,0.00069599,0.0008899135,0.00058608205,0.0048281364,0.009769735,0.004790874,0.0056726225,0.15169188],"study_design_scores_gemma":[0.00020039202,0.0027867686,0.9859368,0.00004519796,0.00017134719,0.00006832933,0.00066417543,0.0013746556,0.0025416014,0.00046561,0.0057202317,0.000024878413],"about_ca_topic_score_codex":0.004955074,"about_ca_topic_score_gemma":0.005518193,"teacher_disagreement_score":0.012584949,"about_ca_system_score_codex":0.0020379198,"about_ca_system_score_gemma":0.0018038711,"threshold_uncertainty_score":0.042100847},"labels":[],"label_agreement":null},{"id":"W7037547335","doi":"","title":"Effect of exhaust gas recirculation on fuel consumption and nitrogen oxides emissions","year":2001,"lang":"en","type":"other","venue":"Library and Archives Canada (Government of Canada)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Exhaust gas recirculation; Nitrogen oxides; Nitrogen; Consumption (sociology); Exhaust gas; Fuel efficiency; Combustion; Air pollution","score_opus":0.004566839285361634,"score_gpt":0.18363880390311726,"score_spread":0.17907196461775562,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037547335","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99266523,0.00030935093,0.0002049122,0.00006614637,0.000016093862,0.00002070205,0.0027649303,0.00010335923,0.0038492244],"genre_scores_gemma":[0.9830725,0.00022967835,0.00033230893,0.00008800243,0.0000030547487,0.000011555216,0.002747319,0.00004192762,0.013473584],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99973816,0.000035966237,0.000008049097,0.00004860842,0.00009362025,0.00007542439],"domain_scores_gemma":[0.9985103,0.00079518824,0.000083478146,0.00006132407,0.0004314125,0.00011834757],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051672745,0.00039436747,0.0004441754,0.00033260428,0.00043112904,0.00065513636,0.0006042827,0.0005393377,0.0054727877],"category_scores_gemma":[0.0010950065,0.00023196448,0.00063890853,0.00051811035,0.00020865198,0.00020828127,0.00016172719,0.0003890937,0.00086523336],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.09851045,0.0045731952,0.48005775,0.000526969,0.0015705442,0.0015471827,0.00051728927,0.08969578,0.16121234,0.00048603502,0.014211843,0.14709061],"study_design_scores_gemma":[0.0002276438,0.0024127024,0.8915461,0.000025754804,0.0005939834,0.00015103254,0.0005832477,0.016280634,0.083843604,0.00012426903,0.004145556,0.00006543062],"about_ca_topic_score_codex":0.6213785,"about_ca_topic_score_gemma":0.69975907,"teacher_disagreement_score":0.6213785,"about_ca_system_score_codex":0.0028365622,"about_ca_system_score_gemma":0.0033267082,"threshold_uncertainty_score":0.7617026},"labels":[],"label_agreement":null},{"id":"W7037649926","doi":"","title":"Esteettömyys Oulun katurakentamisessa","year":2013,"lang":"fi","type":"other","venue":"Theseus (Ammattikorkeakoulujen)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Work (physics); Power (physics); Data collection; Quarter (Canadian coin)","score_opus":0.017384662040950728,"score_gpt":0.2634161449700484,"score_spread":0.24603148292909766,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7037649926","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.059205767,0.011694098,0.012770916,0.015104104,0.006803517,0.00025172246,0.0010495561,0.0011020651,0.89201814],"genre_scores_gemma":[0.08296643,0.006678643,0.010671754,0.002229126,0.0005720323,0.00011373838,0.0008404197,0.00052953325,0.8953982],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989906,0.00008795584,0.000057967798,0.00018068068,0.00046723938,0.00021563012],"domain_scores_gemma":[0.99880624,0.00019381924,0.00011444096,0.0001033604,0.0005218557,0.00026024212],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00093491486,0.0007332873,0.0004884816,0.0009369004,0.0027370749,0.0060606925,0.0008390552,0.0012542102,0.12927146],"category_scores_gemma":[0.0014746485,0.00038226967,0.0006958612,0.0006413244,0.0014332756,0.0029767558,0.004391694,0.0024305568,0.037418157],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009678882,0.0004615496,0.012890877,0.0015676133,0.00006703845,0.0027231462,0.006285055,0.0005556577,0.07595146,0.062440224,0.172864,0.66322553],"study_design_scores_gemma":[0.000009436181,0.00009582664,0.0036290477,0.0001965232,0.000018927707,0.00064090284,0.002330805,0.00020018514,0.0106036095,0.0016969899,0.98054975,0.000027881131],"about_ca_topic_score_codex":0.007158389,"about_ca_topic_score_gemma":0.027055247,"teacher_disagreement_score":0.12927146,"about_ca_system_score_codex":0.0025391271,"about_ca_system_score_gemma":0.0037403887,"threshold_uncertainty_score":0.43245614},"labels":[],"label_agreement":null},{"id":"W7038406346","doi":"","title":"The Impact of the Severe Acute Respiratory Syndrome (SARS) on International Airline Demand in Asia Pacific","year":2008,"lang":"en","type":"other","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Lagging; China; Profitability index; Asia pacific; Southeast asia; Aviation; International airport; Economic impact analysis","score_opus":0.01289202983088589,"score_gpt":0.2903459633596217,"score_spread":0.2774539335287358,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7038406346","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99461395,0.00016852024,0.00018551023,0.00026771874,0.0000058624846,0.000005538761,0.0006160565,0.000003092602,0.0041338345],"genre_scores_gemma":[0.99850225,0.00023116384,0.000069452995,0.000021400507,0.0000068053714,0.00000468964,0.00042306542,0.0000031312095,0.00073802395],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9995307,0.00013342619,0.0000363165,0.00005719341,0.00012589566,0.000116493204],"domain_scores_gemma":[0.99752957,0.0010444481,0.00067752344,0.00007779576,0.00043832103,0.00023234342],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00051224837,0.0002508637,0.00029804168,0.00049573486,0.00026206087,0.0017128487,0.00030333846,0.00040568245,0.0029259904],"category_scores_gemma":[0.002525792,0.00018284873,0.0005139494,0.0012285092,0.00026244717,0.00087480526,0.0006249144,0.00083092705,0.00033954528],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016566053,0.00016378253,0.9644087,0.00009750696,0.00016021969,0.00065360154,0.00082205393,0.021227486,0.0009683954,0.0012554709,0.0009676196,0.009109536],"study_design_scores_gemma":[0.0000040278337,0.000120677694,0.97528744,0.00003322171,0.000045515215,0.00016345562,0.0043191705,0.017494436,0.00040800552,0.0003627615,0.0017432157,0.000018064207],"about_ca_topic_score_codex":0.028218884,"about_ca_topic_score_gemma":0.022973191,"teacher_disagreement_score":0.028218884,"about_ca_system_score_codex":0.001187574,"about_ca_system_score_gemma":0.00079733564,"threshold_uncertainty_score":0.05610931},"labels":[],"label_agreement":null},{"id":"W7038716818","doi":"","title":"Information seeking and sharing among doctoral peers: An exploratory case study in the context of skills","year":2025,"lang":"en","type":"dissertation","venue":"eScholarship@McGill (McGill)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Information seeking; Context (archaeology); Exploratory research; Information sharing; Information seeking behavior","score_opus":0.02312516625090954,"score_gpt":0.28328014641474486,"score_spread":0.2601549801638353,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7038716818","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9902794,0.00017482259,0.00048905634,0.002666396,0.000024431312,0.00008271141,0.000032665313,0.000011104544,0.006239433],"genre_scores_gemma":[0.9961367,0.00021709039,0.00070477556,0.00039762404,0.000018107841,0.000058483827,0.00003058773,0.000009715519,0.0024268539],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9907266,0.006279961,0.000249374,0.0005945447,0.0009877018,0.0011618433],"domain_scores_gemma":[0.9586962,0.02979867,0.0021952002,0.0010056783,0.0017271355,0.0065771905],"candidate_categories":["scholarly_communication"],"consensus_categories":[],"category_scores_codex":[0.008863627,0.00038927715,0.00053264573,0.0021147593,0.017319772,0.006000244,0.0017470489,0.004102683,0.0048472704],"category_scores_gemma":[0.028222438,0.00043506618,0.00042262208,0.0019496517,0.0044327634,0.0057169823,0.006311119,0.0029359844,0.00061533466],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010861213,0.0009316865,0.018366681,0.00010769806,0.000015756763,0.005584788,0.9537319,0.000079351616,0.0005937539,0.0020515998,0.0016798578,0.016748255],"study_design_scores_gemma":[0.00001888671,0.00026368917,0.0050219838,0.000058053127,0.000011643714,0.0014439503,0.9853892,0.00023717512,0.00026309813,0.00071769085,0.006557624,0.000017079165],"about_ca_topic_score_codex":0.00976138,"about_ca_topic_score_gemma":0.017304959,"teacher_disagreement_score":0.9939998,"about_ca_system_score_codex":0.003326406,"about_ca_system_score_gemma":0.0059300177,"threshold_uncertainty_score":0.046875894},"labels":[],"label_agreement":null},{"id":"W7039000422","doi":"","title":"An Intermodel comparison of DDS and Daysim daylight coefficient models","year":2006,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"Natural Resources Canada; National Research Council Canada","keywords":"Daylighting; Daylight; Independence (probability theory); Doors; Correlation coefficient; Standard uncertainty","score_opus":0.01869456226845907,"score_gpt":0.29205139214773507,"score_spread":0.273356829879276,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039000422","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8196828,0.00020824946,0.16797268,0.00023862055,0.00011471172,0.00015031779,0.0027757725,0.002494316,0.006362635],"genre_scores_gemma":[0.9708455,0.000048257727,0.025735535,0.000044539,0.0000074545737,0.000096400785,0.0022162828,0.00025597008,0.00075014535],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989819,0.0003403952,0.00009446846,0.00022169761,0.00026952347,0.00009202103],"domain_scores_gemma":[0.99661463,0.0016000457,0.00019889366,0.0005970846,0.0009128896,0.000076512246],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027876017,0.0007891686,0.00075464475,0.0007145702,0.00044675675,0.00085912447,0.0012507169,0.0005744885,0.0021728633],"category_scores_gemma":[0.0059768385,0.00040931904,0.0009945645,0.00057156157,0.0003222045,0.0012027812,0.0008516951,0.0007647958,0.00036997654],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004379365,0.00013002119,0.005199204,0.00006793844,0.00016699401,0.000026974556,0.000049528804,0.9719454,0.0024672386,0.0010195214,0.00068542303,0.017803913],"study_design_scores_gemma":[0.00003527321,0.00013473576,0.0023169883,0.000004900683,0.00003118406,0.000012493083,0.00003139353,0.99249554,0.0034815602,0.0005848576,0.0008552325,0.000015814408],"about_ca_topic_score_codex":0.018590448,"about_ca_topic_score_gemma":0.013032462,"teacher_disagreement_score":0.018590448,"about_ca_system_score_codex":0.0014464946,"about_ca_system_score_gemma":0.0011258546,"threshold_uncertainty_score":0.036964476},"labels":[],"label_agreement":null},{"id":"W7039684678","doi":"","title":"Morphological Variation in Haitian Creole","year":2019,"lang":"en","type":"report","venue":"IUScholarWorks (Indiana University)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"National Science Foundation","keywords":"Creole language; Possessive; Pronoun; Context (archaeology); Variety (cybernetics); Variation (astronomy); Focus (optics); Tamil","score_opus":0.0235979663796502,"score_gpt":0.2538970342776603,"score_spread":0.2302990678980101,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7039684678","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99099195,0.00015000405,0.00009868665,0.000046491165,0.000005273465,0.000011429475,0.00017716784,0.000007243405,0.008511789],"genre_scores_gemma":[0.99654835,0.0001787066,0.0003217323,0.000045323894,0.0000059985077,0.000012593049,0.0003090354,0.000021304273,0.002556947],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9997484,0.00003773613,0.000015968399,0.00007904406,0.00006205044,0.000056721998],"domain_scores_gemma":[0.9997031,0.000095291485,0.00005337854,0.000030222534,0.000091705966,0.000026251784],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002979664,0.00020037823,0.0001760525,0.0016928698,0.001196044,0.0008321556,0.0003679554,0.0003583304,0.0038316662],"category_scores_gemma":[0.00089124776,0.00016523867,0.0001324253,0.0016585055,0.0012327572,0.00022694815,0.0007613777,0.00025791395,0.0006836692],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0010304117,0.0001716918,0.42298585,0.00047890336,0.00015191754,0.013097885,0.21313158,0.00037321972,0.15118828,0.0073691565,0.0050266264,0.1849945],"study_design_scores_gemma":[0.00001151669,0.00006235514,0.943532,0.00003558358,0.000023207956,0.004743815,0.02358157,0.00017134954,0.0029610258,0.0002301248,0.02460754,0.000039995488],"about_ca_topic_score_codex":0.03646142,"about_ca_topic_score_gemma":0.08333009,"teacher_disagreement_score":0.03646142,"about_ca_system_score_codex":0.00080661575,"about_ca_system_score_gemma":0.000640879,"threshold_uncertainty_score":0.07249838},"labels":[],"label_agreement":null},{"id":"W7065861266","doi":"","title":"Expirations of Pandemic Jobless Programs Caused an Unprecedented Drop in Access to UI","year":2022,"lang":"en","type":"other","venue":"eScholarship (California Digital Library)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Unemployment; Disadvantaged; Population; Pandemic; Recession; Working poor; Economic shortage; Quarter (Canadian coin)","score_opus":0.034349299406695184,"score_gpt":0.2864584597787147,"score_spread":0.2521091603720195,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7065861266","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9411707,0.00063799525,0.0012612037,0.013089847,0.0006545628,0.00022412278,0.0084333485,0.00043026463,0.034098003],"genre_scores_gemma":[0.9801176,0.0004962895,0.0007005115,0.005244095,0.00039052634,0.00017165429,0.00399545,0.00004397668,0.0088399295],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.998362,0.0001402134,0.00007176564,0.00016267024,0.0006122961,0.00065096957],"domain_scores_gemma":[0.9974591,0.0005656427,0.0006680354,0.00016267122,0.00045877026,0.00068580016],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019784858,0.00027476336,0.00028096887,0.00072138954,0.0011061971,0.0018363993,0.0006033257,0.0008681841,0.0042700707],"category_scores_gemma":[0.005850931,0.00029305444,0.00070563855,0.00047785704,0.00057607115,0.000852652,0.0022895017,0.0032507898,0.00040732414],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009886308,0.0009797919,0.49472886,0.0004960287,0.00023553125,0.0014600618,0.0044390024,0.0059605464,0.0063689826,0.010190804,0.15357473,0.32057694],"study_design_scores_gemma":[0.000093278206,0.0006833885,0.9237491,0.0001820803,0.000065403234,0.00025507688,0.003185012,0.0028042076,0.0021918227,0.00059406017,0.06614269,0.000053862066],"about_ca_topic_score_codex":0.06279826,"about_ca_topic_score_gemma":0.07970138,"teacher_disagreement_score":0.06279826,"about_ca_system_score_codex":0.0026107323,"about_ca_system_score_gemma":0.003213303,"threshold_uncertainty_score":0.12486547},"labels":[],"label_agreement":null},{"id":"W7067012245","doi":"","title":"Le népotisme filial dans un groupe captif de macaques crabiers (Macaca fascicularis)","year":2005,"lang":"fr","type":"other","venue":"Open MIND","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Qualitative research; Perspective (graphical); Spouse; Population","score_opus":0.01397167192679916,"score_gpt":0.2733717581017227,"score_spread":0.25940008617492355,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7067012245","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.955188,0.0070566204,0.0076971506,0.0018553982,0.00009128184,0.000040976545,0.0078314645,0.00026188086,0.019977242],"genre_scores_gemma":[0.96666783,0.0023259327,0.009182869,0.00010680848,0.000046855454,0.000053005653,0.0045484845,0.00008483642,0.016983321],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.99965036,0.00008436978,0.00002267527,0.00011011725,0.00008468893,0.00004779338],"domain_scores_gemma":[0.9981353,0.0007837318,0.0003179099,0.00014958896,0.00048651788,0.00012687329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008470382,0.00033828276,0.00030323752,0.0046080495,0.0020160829,0.0019312524,0.00046586653,0.0005637409,0.0052868105],"category_scores_gemma":[0.005255874,0.00011187221,0.0003528188,0.0029509992,0.0008165657,0.0012111333,0.0010161544,0.0005599385,0.00049068825],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009166598,0.00012651864,0.44344702,0.00080867705,0.000533403,0.0023889886,0.028855495,0.0019724576,0.02126667,0.016796397,0.015579617,0.4673082],"study_design_scores_gemma":[0.000024836292,0.00008416316,0.88253164,0.0002106837,0.0002308228,0.0026683065,0.008495467,0.0025512564,0.004283975,0.003650972,0.09520741,0.000060500035],"about_ca_topic_score_codex":0.33081728,"about_ca_topic_score_gemma":0.41299593,"teacher_disagreement_score":0.33081728,"about_ca_system_score_codex":0.0025527913,"about_ca_system_score_gemma":0.0015185146,"threshold_uncertainty_score":0.6577834},"labels":[],"label_agreement":null},{"id":"W7070919841","doi":"","title":"Protein name tagging","year":2000,"lang":"en","type":"article","venue":"NPARC","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"","score_opus":0.01013900926247054,"score_gpt":0.2459381621994657,"score_spread":0.23579915293699516,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7070919841","genre_codex":"dataset","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005730799,0.0021698135,0.008538576,0.006161742,0.019236682,0.0005433537,0.69175905,0.019658485,0.24620152],"genre_scores_gemma":[0.01630179,0.0024315033,0.0149100665,0.002323559,0.0033125097,0.00029730634,0.7409097,0.0067253294,0.21278813],"study_design_codex":"not_applicable","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9982127,0.00014893337,0.00027077924,0.0004756577,0.0006869312,0.00020491192],"domain_scores_gemma":[0.99003625,0.0015505272,0.0007147211,0.0025123272,0.004171427,0.0010146549],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0020323948,0.0017344069,0.0017079017,0.008713081,0.001692935,0.004041848,0.0015098613,0.0025582986,0.4087457],"category_scores_gemma":[0.008897442,0.000888902,0.0010557243,0.008512732,0.00046042353,0.0038750155,0.0020295805,0.0015179805,0.43764246],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000317151,0.000070383394,0.00065470434,0.00079095585,0.000017015733,0.00014293833,0.000026151834,0.000064410895,0.0059580794,0.0016973069,0.954759,0.035501998],"study_design_scores_gemma":[0.000062564926,0.00002783404,0.0020706982,0.00016098349,0.00003101134,0.00020518397,0.00003184175,0.00026388263,0.0051681874,0.0011185823,0.9908283,0.00003097342],"about_ca_topic_score_codex":0.0033601471,"about_ca_topic_score_gemma":0.005513063,"teacher_disagreement_score":0.4087457,"about_ca_system_score_codex":0.0015440851,"about_ca_system_score_gemma":0.002968457,"threshold_uncertainty_score":0.8433525},"labels":[],"label_agreement":null},{"id":"W7073773086","doi":"","title":"PreBIND and Textomy – mining the biomedical literature for protein-protein interactions using a support vector machine","year":2003,"lang":"en","type":"article","venue":"TSpace","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Alpha Cancer Technologies; National Research Council Canada; Lunenfeld-Tanenbaum Research Institute","funders":"Canadian Institutes of Health Research; Directorate for Biological Sciences; Genome Canada","keywords":"Support vector machine; Feature (linguistics); Identification (biology); Feature selection","score_opus":0.02649095008400263,"score_gpt":0.34502766077850916,"score_spread":0.3185367106945065,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7073773086","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.30211934,0.03314395,0.46547186,0.005545646,0.00094877847,0.0039034497,0.12983009,0.045650993,0.013385891],"genre_scores_gemma":[0.2908183,0.0059819473,0.5969889,0.0006854841,0.00041017783,0.0018498966,0.094804496,0.00049376395,0.007967058],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997766,0.00039641836,0.0005302757,0.0005602665,0.0006361791,0.00011080045],"domain_scores_gemma":[0.99294925,0.0033450327,0.0010485497,0.0005839279,0.0018160405,0.00025723278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003030881,0.0011128733,0.0012067162,0.015766695,0.0006664349,0.0021103404,0.0013519716,0.0008568031,0.0076265894],"category_scores_gemma":[0.009907479,0.000366653,0.001239571,0.0070842425,0.00055618933,0.00197593,0.0013261215,0.00076813664,0.004204184],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00087606674,0.00034714883,0.018371012,0.0061711073,0.00038365123,0.0011191927,0.0005740114,0.005645162,0.04007928,0.0028001936,0.043360036,0.88027316],"study_design_scores_gemma":[0.00057288783,0.0026577835,0.089393124,0.0021043455,0.0011219183,0.0078066443,0.0024703068,0.31162405,0.19767916,0.035976414,0.3481727,0.0004207077],"about_ca_topic_score_codex":0.0021236835,"about_ca_topic_score_gemma":0.0033473806,"teacher_disagreement_score":0.015766695,"about_ca_system_score_codex":0.0007671319,"about_ca_system_score_gemma":0.0027717047,"threshold_uncertainty_score":0.02551347},"labels":[],"label_agreement":null},{"id":"W7089445583","doi":"10.15468/dl.5d7th2","title":"Occurrence Download","year":2025,"lang":"en","type":"dataset","venue":"Global Biodiversity Information Facility","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Download; Matching (statistics); Range (aeronautics); Set (abstract data type); Data set","score_opus":0.011426349904950839,"score_gpt":0.2390157115319853,"score_spread":0.22758936162703447,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7089445583","genre_codex":"dataset","genre_gemma":"dataset","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"dataset","genre_consensus":"dataset","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00008027137,0.000046501347,0.00006535097,0.000054171396,0.000013771463,0.0000081407225,0.99849164,0.00052945875,0.000710668],"genre_scores_gemma":[0.00018387765,0.000038017817,0.000228207,0.00004090348,0.0000028933755,0.000038261038,0.9989986,0.00010035647,0.00036892173],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9989373,0.00014890224,0.00014681954,0.0003604024,0.00025593746,0.0001506908],"domain_scores_gemma":[0.9978654,0.00063547323,0.00019825132,0.00054497,0.00048353113,0.00027232795],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0010488399,0.0022874332,0.0015580218,0.0052618203,0.0010401286,0.0025016502,0.0029167647,0.002258554,0.10260222],"category_scores_gemma":[0.0056931474,0.0008670651,0.0012216129,0.009115531,0.00048189404,0.0022370918,0.00244419,0.0019962585,0.15023647],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000041724146,0.000018615472,0.0004935282,0.000647785,0.000020576168,0.000026134121,0.000027531203,0.00020557,0.0001639533,0.00048757056,0.995904,0.0019629418],"study_design_scores_gemma":[0.00009342136,0.0000109767225,0.0020898979,0.00020574164,0.000018474202,0.00006410159,0.00008092644,0.00030609025,0.0002856658,0.0009833355,0.99584156,0.000019634279],"about_ca_topic_score_codex":0.019213995,"about_ca_topic_score_gemma":0.03290722,"teacher_disagreement_score":0.89739776,"about_ca_system_score_codex":0.001731369,"about_ca_system_score_gemma":0.0024474766,"threshold_uncertainty_score":0.3432386},"labels":[],"label_agreement":null},{"id":"W7094145776","doi":"","title":"Erratum: Correction of author name and credentials and errors in text: Effectiveness of the CANRISK tool in the identification of dysglycemia in First Nations and MÃ©tis in Canada","year":2018,"lang":"en","type":"article","venue":"PubMed Central","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Identification (biology); MEDLINE","score_opus":0.008459775617100618,"score_gpt":0.23351801600781194,"score_spread":0.22505824039071132,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7094145776","genre_codex":"editorial","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.001188238,0.0010557644,0.002606782,0.07731522,0.90806556,0.00012216442,0.0058461777,0.0011273784,0.0026725885],"genre_scores_gemma":[0.073238276,0.012188192,0.051593587,0.2398323,0.23008649,0.0011439298,0.025299897,0.012546885,0.35407045],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9853086,0.0017295611,0.00424145,0.0012634778,0.0064237467,0.0010332006],"domain_scores_gemma":[0.8253401,0.0413296,0.006781927,0.008174657,0.11511616,0.0032575978],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009146776,0.0017869881,0.0019515064,0.007914544,0.0045641484,0.0053652194,0.0037599157,0.0064221243,0.060409393],"category_scores_gemma":[0.16420935,0.0015880768,0.0026103128,0.005568578,0.0025456692,0.0024336097,0.0031411063,0.007626805,0.026754316],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008635011,0.000017436674,0.00032498405,0.00022857764,0.000032494452,0.00048984186,0.00009380494,0.00005783257,0.00012366829,0.00034852634,0.9922587,0.005937822],"study_design_scores_gemma":[0.00012742713,0.000056783872,0.002607839,0.0011123857,0.00023269154,0.0014227094,0.0006326162,0.0008264121,0.0017347494,0.0013141601,0.9898073,0.00012504998],"about_ca_topic_score_codex":0.08913393,"about_ca_topic_score_gemma":0.08515752,"teacher_disagreement_score":0.9108661,"about_ca_system_score_codex":0.006526171,"about_ca_system_score_gemma":0.018285774,"threshold_uncertainty_score":0.20208955},"labels":[],"label_agreement":null},{"id":"W7095081523","doi":"","title":"arrangement with Health Canada’s Health Protection Branch.","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Health protection; Public health; Government (linguistics); Data Protection Act 1998; Health policy","score_opus":0.0157990974863538,"score_gpt":0.2436678579494508,"score_spread":0.227868760463097,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7095081523","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.005010529,0.0031963517,0.0056710765,0.3206166,0.018898137,0.0012304819,0.032075666,0.002654116,0.6106471],"genre_scores_gemma":[0.009085381,0.00036845548,0.0014436378,0.010542936,0.0006144882,0.00008588372,0.003342025,0.000245986,0.97427124],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99516016,0.00033123346,0.00009643452,0.00071541994,0.0025431658,0.0011536755],"domain_scores_gemma":[0.9764227,0.001094173,0.00036109224,0.0010298713,0.011190673,0.009901451],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0041449917,0.00064082915,0.00076255033,0.0017890254,0.0064559155,0.0045772153,0.001703359,0.0040149414,0.40816402],"category_scores_gemma":[0.0079752635,0.0004498857,0.000615191,0.0022851143,0.0012755995,0.0019118387,0.0022823329,0.0031367152,0.11511603],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010457643,0.00003903452,0.0009420023,0.000024953946,0.000010332086,0.00007299977,0.000060822822,0.00003746406,0.00023763232,0.0042340388,0.97540593,0.018830238],"study_design_scores_gemma":[0.00003898397,0.000022139582,0.0028264076,0.000034261233,0.00000927981,0.00006066821,0.00026129565,0.00021858799,0.00024772124,0.0011050688,0.99516135,0.000014127179],"about_ca_topic_score_codex":0.6120901,"about_ca_topic_score_gemma":0.7360679,"teacher_disagreement_score":0.6120901,"about_ca_system_score_codex":0.013988951,"about_ca_system_score_gemma":0.08497472,"threshold_uncertainty_score":0.8441822},"labels":[],"label_agreement":null},{"id":"W7097263554","doi":"","title":"Ontology-based Representation and Analysis of Vaccination Informed Consent","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Informed consent; Vaccination; Ontology; Government (linguistics); Representation (politics)","score_opus":0.03224439955716472,"score_gpt":0.3335558764028841,"score_spread":0.30131147684571935,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7097263554","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023249242,0.0003507993,0.9516896,0.0022473251,0.000100877696,0.000999778,0.0032520983,0.0025018558,0.015608464],"genre_scores_gemma":[0.1912223,0.0005912196,0.79461354,0.0004055704,0.000043519216,0.00064501667,0.008552807,0.00033425723,0.0035917417],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99096596,0.0033946114,0.0015221654,0.00085599197,0.0027921714,0.00046919342],"domain_scores_gemma":[0.98688716,0.005934265,0.0013160086,0.0026870652,0.0028349657,0.00034060152],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009394481,0.0005629323,0.00051521434,0.0064177923,0.0017442901,0.005568626,0.0014446649,0.0013598748,0.0023115857],"category_scores_gemma":[0.019333474,0.00053837034,0.0021828804,0.0046732463,0.0017765514,0.006470325,0.0036135232,0.001854073,0.00046148992],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00016533927,0.00033194825,0.009524446,0.0007138561,0.00016606826,0.0012658983,0.006795814,0.044811007,0.0064149178,0.7448632,0.012940429,0.17200708],"study_design_scores_gemma":[0.00006129118,0.000080988044,0.004617791,0.00094424636,0.00019060972,0.00090365886,0.0041124867,0.39056966,0.0129498495,0.34681663,0.23861119,0.00014166102],"about_ca_topic_score_codex":0.030609805,"about_ca_topic_score_gemma":0.02605709,"teacher_disagreement_score":0.030609805,"about_ca_system_score_codex":0.00438044,"about_ca_system_score_gemma":0.0069918735,"threshold_uncertainty_score":0.060863316},"labels":[],"label_agreement":null},{"id":"W7097596706","doi":"","title":"InternationalJterna","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Knowledge base; XML; Parsing; Decision rule; Decision table; Base (topology); Decision support system; Representation (politics)","score_opus":0.009881245535212914,"score_gpt":0.27395405284282937,"score_spread":0.26407280730761645,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7097596706","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010274512,0.006720291,0.009020638,0.0057236166,0.0028344525,0.00020068689,0.0052847415,0.0021340488,0.9578069],"genre_scores_gemma":[0.02885816,0.004365479,0.0105700325,0.0012946507,0.00022740489,0.00015507963,0.0055615045,0.0007496283,0.94821805],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99875605,0.0001707304,0.000113276226,0.0003263234,0.0004426168,0.00019097589],"domain_scores_gemma":[0.99921155,0.0001375704,0.000111085916,0.00016795499,0.00023832796,0.00013342673],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0012122777,0.0007778521,0.00051156693,0.002937512,0.0020397264,0.005135736,0.0012224801,0.0015833307,0.32642257],"category_scores_gemma":[0.00232932,0.0004158534,0.0006733471,0.0026664268,0.0007951205,0.0023460435,0.0034605693,0.0016430762,0.16022877],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00040634393,0.00021022523,0.0045565097,0.0008707736,0.000049056867,0.0017291114,0.0012117648,0.00089104107,0.0058337147,0.13218616,0.21404272,0.63801265],"study_design_scores_gemma":[0.000011201659,0.000016889999,0.0013458006,0.0001242358,0.000009137098,0.0002911518,0.00014174242,0.00019849339,0.00078903424,0.0025834357,0.99447906,0.000009771724],"about_ca_topic_score_codex":0.010659377,"about_ca_topic_score_gemma":0.016164616,"teacher_disagreement_score":0.6735774,"about_ca_system_score_codex":0.0026398946,"about_ca_system_score_gemma":0.0028870364,"threshold_uncertainty_score":0.96077645},"labels":[],"label_agreement":null},{"id":"W7097725207","doi":"","title":"Sponsored by the AMIA Formal Biomedical Knowledge Representation Special Interest Group","year":2004,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Knowledge representation and reasoning; Special Interest Group; Representation (politics); Ninth; Set (abstract data type); Domain (mathematical analysis); Health informatics","score_opus":0.025691084605745002,"score_gpt":0.3030643612677903,"score_spread":0.2773732766620453,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7097725207","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00597225,0.02065899,0.1975022,0.14139822,0.042014115,0.0029107807,0.017184954,0.01406625,0.55829227],"genre_scores_gemma":[0.022181755,0.01212695,0.09945673,0.0070326678,0.011699731,0.0012722461,0.035302326,0.004630518,0.8062971],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99697113,0.0006901421,0.0002585339,0.00055226794,0.0012193861,0.00030860826],"domain_scores_gemma":[0.98660046,0.0023597397,0.0005998872,0.0014500985,0.0054351115,0.0035546992],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.009655537,0.0011333418,0.0013505219,0.0023499953,0.0013895895,0.0063684625,0.0020877537,0.003384602,0.2511229],"category_scores_gemma":[0.010056526,0.0006096689,0.0012499611,0.002726061,0.0007783474,0.0058673876,0.0044226404,0.0029273513,0.14588608],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00018916655,0.000098131946,0.0002551172,0.00027764798,0.00001698529,0.00007988586,0.00010534691,0.00037510818,0.0023691554,0.018385563,0.79587406,0.18197381],"study_design_scores_gemma":[0.00004408399,0.000045055967,0.00050259486,0.00012024725,0.000014084811,0.00011989422,0.00006538265,0.001850524,0.0005631521,0.0073090903,0.98934877,0.000017297258],"about_ca_topic_score_codex":0.004477663,"about_ca_topic_score_gemma":0.0036173193,"teacher_disagreement_score":0.7488771,"about_ca_system_score_codex":0.0021636104,"about_ca_system_score_gemma":0.005460393,"threshold_uncertainty_score":0.8400898},"labels":[],"label_agreement":null},{"id":"W7098216775","doi":"","title":"The Ins and Outs of Poverty in Advanced Economies: Government Policy and Poverty Dynamics in","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Poverty; Chronic poverty; Government (linguistics); Public policy; Culture of poverty; Basic needs","score_opus":0.003831441693398207,"score_gpt":0.22942459388439393,"score_spread":0.2255931521909957,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098216775","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9903447,0.00040499296,0.0003691738,0.0021727239,0.0000043925547,0.0000052070927,0.00015431,0.000005533932,0.006539002],"genre_scores_gemma":[0.99955744,0.00017494425,0.00006600422,0.000028775994,0.000001908996,0.000002767929,0.00003117724,0.0000013070182,0.00013563402],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9993932,0.00021758956,0.000031900385,0.000049776998,0.000098229095,0.00020912457],"domain_scores_gemma":[0.9978727,0.0007478008,0.00058368006,0.0001222907,0.00036767332,0.00030578463],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0014249899,0.000076538934,0.00026446444,0.0020963007,0.0011328188,0.002089846,0.00022720904,0.00044703326,0.0014922766],"category_scores_gemma":[0.004416879,0.00008472568,0.00018585147,0.003583519,0.0025695907,0.0037848498,0.0019013189,0.0006202772,0.00006718893],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00078308233,0.00036632706,0.68409264,0.0001610041,0.00024690083,0.0005441501,0.02589501,0.012518019,0.0014014803,0.16454348,0.0031225765,0.10632531],"study_design_scores_gemma":[0.0000141401615,0.00009857074,0.91928005,0.000104893785,0.000030869345,0.00010174394,0.025362207,0.0030533853,0.00042705273,0.041699223,0.009799585,0.000028217568],"about_ca_topic_score_codex":0.025641033,"about_ca_topic_score_gemma":0.037593443,"teacher_disagreement_score":0.025641033,"about_ca_system_score_codex":0.0033663136,"about_ca_system_score_gemma":0.0026685288,"threshold_uncertainty_score":0.05098355},"labels":[],"label_agreement":null},{"id":"W7098322400","doi":"","title":"1RECOVERY FROM DEPRESSION:","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Debt; Exchange rate; Depression (economics); Foreign exchange; Dutch disease; External debt","score_opus":0.010195327061006197,"score_gpt":0.24031650967697887,"score_spread":0.23012118261597267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098322400","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16421214,0.041205086,0.014773716,0.327185,0.0076817474,0.00021598497,0.007636562,0.00086206256,0.43622765],"genre_scores_gemma":[0.809385,0.026841948,0.00898159,0.030874116,0.0024724896,0.00018917782,0.009805282,0.0005388028,0.11091162],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99838495,0.0005004167,0.00009885107,0.00029549949,0.00041455685,0.00030575463],"domain_scores_gemma":[0.99649185,0.00084739435,0.00064771844,0.00034403216,0.0008616904,0.00080728263],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0023924352,0.00034632575,0.0003254591,0.0016142305,0.0035682383,0.0059848414,0.0008874502,0.0015944784,0.019865774],"category_scores_gemma":[0.009929906,0.00017542347,0.00038726119,0.0032022335,0.0027400362,0.009792878,0.005071467,0.0024190661,0.005312621],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023163228,0.00006972072,0.036653318,0.0010253796,0.000052731913,0.0028730866,0.05247799,0.0002091151,0.0008801157,0.15659548,0.40507033,0.3438611],"study_design_scores_gemma":[0.000009319605,0.000035775072,0.029989975,0.0010407873,0.000025100955,0.001863064,0.03647065,0.0003908987,0.0004655945,0.03688961,0.8927821,0.000037145768],"about_ca_topic_score_codex":0.013386877,"about_ca_topic_score_gemma":0.024545288,"teacher_disagreement_score":0.98013425,"about_ca_system_score_codex":0.004149553,"about_ca_system_score_gemma":0.0042497576,"threshold_uncertainty_score":0.06645763},"labels":[],"label_agreement":null},{"id":"W7098486467","doi":"","title":"Carcass composition and yield of 1957 versus 2001 broilers when fed representative 1957 and 2001 broiler diets. Poult. Sci","year":2003,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Broiler; Yield (engineering); Carcass weight; Composition (language); Animal production","score_opus":0.04562938351894068,"score_gpt":0.3060050157904674,"score_spread":0.2603756322715267,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098486467","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9992111,0.00003644358,0.000025580004,0.0000046076575,9.258417e-7,0.000002208955,0.000473468,0.0000022560703,0.00024336451],"genre_scores_gemma":[0.99556476,0.00006105296,0.00016432065,0.000014860905,8.6124504e-7,0.0000059296835,0.0025978768,0.0000032568705,0.0015871276],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99985504,0.000010592178,0.000012997624,0.000043146505,0.00005483394,0.000023412664],"domain_scores_gemma":[0.999496,0.000047967216,0.00015609678,0.000031203577,0.00017163064,0.00009720261],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00023572237,0.00019461838,0.00019129975,0.0008063887,0.00033540162,0.00030764772,0.00017518793,0.00022859426,0.0008164568],"category_scores_gemma":[0.00038940724,0.00011555581,0.00015895191,0.00045096528,0.00028558844,0.00015188682,0.00016129045,0.00018341134,0.00023504886],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.005854139,0.00041585343,0.6734367,0.00011698583,0.00041651388,0.00038169627,0.0007320807,0.000579688,0.3042354,0.00016858803,0.00075875607,0.012903671],"study_design_scores_gemma":[0.0000033905878,0.00019449797,0.9950191,0.0000020602981,0.000023445109,0.00005127822,0.00016610105,0.000096137504,0.0041958834,0.000003710857,0.000240634,0.0000037265465],"about_ca_topic_score_codex":0.072057776,"about_ca_topic_score_gemma":0.19901691,"teacher_disagreement_score":0.072057776,"about_ca_system_score_codex":0.0013632515,"about_ca_system_score_gemma":0.0002700666,"threshold_uncertainty_score":0.14327675},"labels":[],"label_agreement":null},{"id":"W7098592692","doi":"","title":"Saskatoon By","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Copying; Permission; Work (physics); Head (geology)","score_opus":0.0229501904757762,"score_gpt":0.26968446687491476,"score_spread":0.24673427639913856,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098592692","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0005275006,0.0010987436,0.0016163576,0.0034467597,0.0017042378,0.00010877085,0.025534218,0.0034493804,0.962514],"genre_scores_gemma":[0.0017037173,0.00078432896,0.0008970629,0.0007522729,0.000063799285,0.00005677966,0.00653271,0.00097216625,0.9882371],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99851876,0.00018207253,0.000119502925,0.00037609323,0.00058461603,0.00021901757],"domain_scores_gemma":[0.9958978,0.00051267666,0.00020901833,0.0013272457,0.001331861,0.000721322],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0014636633,0.0013919034,0.0009799387,0.002443796,0.0022476516,0.009420106,0.0014804052,0.0021921135,0.8874138],"category_scores_gemma":[0.004583483,0.00089217006,0.00082266226,0.0064551383,0.0012701192,0.005100205,0.0046179662,0.0027199646,0.8525502],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055839693,0.000034140863,0.00035741596,0.0001361237,0.000008777419,0.00011969168,0.000089492874,0.00007666522,0.00059476297,0.012431808,0.91445816,0.07163707],"study_design_scores_gemma":[0.000005910705,0.0000034665522,0.0002613327,0.00005831615,0.0000016533429,0.000035516045,0.00004890173,0.00003338894,0.0000779524,0.00069019373,0.9987777,0.0000056573535],"about_ca_topic_score_codex":0.055871714,"about_ca_topic_score_gemma":0.093126014,"teacher_disagreement_score":0.1125862,"about_ca_system_score_codex":0.0044664843,"about_ca_system_score_gemma":0.01224624,"threshold_uncertainty_score":0.16059053},"labels":[],"label_agreement":null},{"id":"W7098662582","doi":"","title":"1 Estimating the Dispersion Parameter of the Negative Binomial Distribution for Analyzing Crash Data Using a Bootstrapped Maximum Likelihood Method","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Negative binomial distribution; Maximum likelihood; Dispersion (optics); Monte Carlo method; Maximum likelihood sequence estimation; Resampling; Estimation theory; Binomial distribution; Moment (physics); Quasi-likelihood","score_opus":0.03910310194451424,"score_gpt":0.3301685947530795,"score_spread":0.2910654928085653,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098662582","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010382898,0.00009572451,0.98875296,0.000073364085,0.000011719282,0.00004636357,0.000072097704,0.000264263,0.00030052205],"genre_scores_gemma":[0.25805318,0.00023576179,0.73976,0.000104647304,0.000086769476,0.00038623047,0.0006522735,0.00013718853,0.00058398425],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9915206,0.005979217,0.00037851796,0.00064438756,0.0013503479,0.00012709938],"domain_scores_gemma":[0.9544077,0.038401976,0.002106472,0.0023589653,0.0025540208,0.00017093505],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01153185,0.0006483251,0.00096489093,0.003154984,0.00063173217,0.0011215887,0.0017375842,0.0013521346,0.0018787296],"category_scores_gemma":[0.07792203,0.0004762886,0.00092755764,0.0020662113,0.000845142,0.0022984322,0.001099752,0.0011809883,0.00082576415],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0003848469,0.00020736018,0.055431537,0.0004950206,0.0004041007,0.00070824823,0.00068629865,0.3349871,0.0118956305,0.035849318,0.0038624904,0.55508804],"study_design_scores_gemma":[0.000036953683,0.00006246299,0.00699501,0.00007879844,0.00003555001,0.00048020075,0.00009940596,0.9619701,0.003106792,0.024852768,0.0022271548,0.000054802298],"about_ca_topic_score_codex":0.0022834202,"about_ca_topic_score_gemma":0.0018848653,"teacher_disagreement_score":0.01153185,"about_ca_system_score_codex":0.00061596895,"about_ca_system_score_gemma":0.0008961696,"threshold_uncertainty_score":0.060986936},"labels":[],"label_agreement":null},{"id":"W7098767968","doi":"","title":"PRINCIPAL AUTHORS","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Principal (computer security); Permission; Subject (documents); Conjunction (astronomy); MEDLINE","score_opus":0.05068957465889182,"score_gpt":0.2712448302958639,"score_spread":0.22055525563697206,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098767968","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0061129797,0.0071563767,0.01477635,0.12587625,0.088332385,0.0016667813,0.022837643,0.004440246,0.728801],"genre_scores_gemma":[0.02408733,0.0032645338,0.005040427,0.021053603,0.00638056,0.00090641074,0.007870191,0.0012536179,0.9301434],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9959229,0.0006229601,0.000345166,0.00084040646,0.0017809346,0.0004875786],"domain_scores_gemma":[0.97912484,0.0014630448,0.0008370138,0.0018818942,0.012039822,0.0046533286],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0041310205,0.00071593176,0.0007782799,0.0020378383,0.0021099949,0.005610095,0.0016193427,0.0022916486,0.55161893],"category_scores_gemma":[0.030627234,0.0003977758,0.0005328575,0.0017767308,0.000652805,0.0025053953,0.0027227756,0.0025483402,0.39135036],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009259344,0.00002230555,0.0008478623,0.00021833285,0.0000072595994,0.00022846874,0.0003027744,0.00005204945,0.00023527956,0.007298889,0.91100484,0.07968935],"study_design_scores_gemma":[0.00001330854,0.000013401157,0.00043464123,0.00013489535,0.000005449078,0.00027176813,0.00024238549,0.000034965527,0.0001380652,0.0015049921,0.9971998,0.000006385982],"about_ca_topic_score_codex":0.0017183308,"about_ca_topic_score_gemma":0.0023081407,"teacher_disagreement_score":0.44838107,"about_ca_system_score_codex":0.0031824396,"about_ca_system_score_gemma":0.009128052,"threshold_uncertainty_score":0.6395612},"labels":[],"label_agreement":null},{"id":"W7098772644","doi":"","title":"Between Women and Men, and Providing a Framework for Accommodation Requests.","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Charter; Publicity; Government (linguistics); Legislation; Legislature; Neutrality; Human rights; Order (exchange)","score_opus":0.02079553527678049,"score_gpt":0.29230195603747094,"score_spread":0.2715064207606904,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098772644","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0049098167,0.0008126874,0.0069449916,0.038442783,0.0015949658,0.00026846767,0.00033766872,0.00014311656,0.9465456],"genre_scores_gemma":[0.13425577,0.0005563015,0.0031516145,0.02162067,0.0006275488,0.00039389695,0.00026038935,0.0001582696,0.83897555],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9876331,0.004506572,0.00037518758,0.0013298586,0.0031702737,0.0029849526],"domain_scores_gemma":[0.9952062,0.001106372,0.0003447321,0.0005973059,0.0014449764,0.0013003653],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006531686,0.00065385073,0.0004719131,0.0013990293,0.018870808,0.0101692015,0.0023554917,0.005808672,0.11148723],"category_scores_gemma":[0.012159387,0.00051454373,0.00050287234,0.0014185463,0.010161486,0.0061968416,0.009838562,0.0035589628,0.03654739],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000048642385,0.000047349255,0.0019425501,0.00007412757,0.0000060580924,0.00066444866,0.037611026,0.000094720956,0.00051540136,0.6208897,0.28713712,0.050968815],"study_design_scores_gemma":[0.0000067741967,0.000010786883,0.0007891182,0.00008908783,0.000001994482,0.00014440068,0.016697232,0.00004063421,0.00009371855,0.010518091,0.9715893,0.000018895422],"about_ca_topic_score_codex":0.14823887,"about_ca_topic_score_gemma":0.26642734,"teacher_disagreement_score":0.14823887,"about_ca_system_score_codex":0.011864335,"about_ca_system_score_gemma":0.023495162,"threshold_uncertainty_score":0.372962},"labels":[],"label_agreement":null},{"id":"W7098815877","doi":"","title":"PERSPECTIVES","year":2010,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"","score_opus":0.0059024367226465065,"score_gpt":0.2619641033708921,"score_spread":0.2560616666482456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098815877","genre_codex":"commentary","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0068167597,0.005866056,0.0013968879,0.7381474,0.0037435144,0.00005490867,0.00084242044,0.000059800077,0.2430723],"genre_scores_gemma":[0.27562645,0.016926603,0.004203122,0.40863118,0.0033377141,0.00020059737,0.0011908719,0.00026793557,0.28961545],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9857998,0.0023064315,0.00023056036,0.0014968299,0.005251985,0.0049145226],"domain_scores_gemma":[0.98376274,0.0030226617,0.00048846967,0.000492495,0.0075181113,0.0047156275],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.009891332,0.0007769224,0.0005690728,0.0022868575,0.0203433,0.018740501,0.0033479088,0.009390816,0.057005152],"category_scores_gemma":[0.017923923,0.00027673962,0.00095611945,0.0035830783,0.014599451,0.0054283924,0.0060220715,0.009295906,0.004150414],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008629855,0.000058841517,0.004850337,0.0003585671,0.000030232282,0.00082928065,0.031008074,0.00021953444,0.0005744013,0.37243202,0.53939617,0.05015618],"study_design_scores_gemma":[0.000012381921,0.0000072486428,0.0020684616,0.00040623447,0.000017305656,0.00017562271,0.03822261,0.000044188928,0.00016969617,0.019277606,0.9395712,0.00002757646],"about_ca_topic_score_codex":0.9278694,"about_ca_topic_score_gemma":0.9623071,"teacher_disagreement_score":0.94299483,"about_ca_system_score_codex":0.0736826,"about_ca_system_score_gemma":0.16974719,"threshold_uncertainty_score":0.5346072},"labels":[],"label_agreement":null},{"id":"W7098996332","doi":"","title":"Incident Detection on an Arterial Roadway","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Section (typography); Intersection (aeronautics); Boulevard; Incident report","score_opus":0.012843960491799679,"score_gpt":0.28545861486019597,"score_spread":0.27261465436839627,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7098996332","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.67973,0.00076556817,0.2730279,0.0012830406,0.00052221277,0.0016092949,0.009455123,0.008469555,0.025137287],"genre_scores_gemma":[0.82471025,0.0005382434,0.16109653,0.0001728791,0.00008515475,0.00028191786,0.007876403,0.0002182135,0.0050204718],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9925021,0.002094254,0.0006267133,0.0011031202,0.0031207874,0.00055293203],"domain_scores_gemma":[0.9742619,0.012965143,0.0025220963,0.0021122983,0.0072592148,0.0008794333],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0042814845,0.00091463985,0.0007186322,0.004803423,0.0013268926,0.0025718524,0.0012448457,0.00087428064,0.0032755628],"category_scores_gemma":[0.021294542,0.00029846653,0.0008812777,0.0032663406,0.00055477646,0.002039771,0.0019059187,0.0009911121,0.0012394714],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0021524138,0.0007539059,0.25118452,0.001879149,0.00059040764,0.0011275458,0.00600505,0.0304497,0.010183682,0.0056436528,0.02370377,0.6663262],"study_design_scores_gemma":[0.00037836394,0.0018540826,0.34200147,0.00084211526,0.0014128028,0.0028431609,0.02796945,0.4036052,0.08424286,0.013734448,0.12061568,0.0005002851],"about_ca_topic_score_codex":0.01498962,"about_ca_topic_score_gemma":0.019257426,"teacher_disagreement_score":0.01498962,"about_ca_system_score_codex":0.0012103267,"about_ca_system_score_gemma":0.0024782903,"threshold_uncertainty_score":0.029804707},"labels":[],"label_agreement":null},{"id":"W7099014975","doi":"","title":"The Use and Misuse of Prediction","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Typology; Abandonment (legal); Commission; Homogeneous; Work (physics); False positive paradox; Empirical research","score_opus":0.06023033799604226,"score_gpt":0.2874990963985307,"score_spread":0.22726875840248845,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099014975","genre_codex":"review","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06997193,0.31835526,0.17939082,0.2611374,0.010964098,0.00068836764,0.0035053305,0.0013426922,0.15464415],"genre_scores_gemma":[0.8670966,0.07068116,0.03078038,0.020690985,0.0052750893,0.00053221156,0.001098892,0.00023723855,0.0036074708],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.90321857,0.06469987,0.0040209056,0.008536222,0.01811197,0.0014125372],"domain_scores_gemma":[0.5466437,0.38536212,0.023034176,0.022873241,0.019928334,0.0021584365],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.10752045,0.001419172,0.001529187,0.0061954586,0.0012376934,0.007212751,0.0064514363,0.0032106675,0.0037974075],"category_scores_gemma":[0.27473527,0.00081979804,0.0015724001,0.004973344,0.015695978,0.015583406,0.0057075904,0.006969545,0.0017178702],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00042974157,0.000106724736,0.06022259,0.0037750956,0.00073435396,0.00033525872,0.006424882,0.004202899,0.000069778784,0.1774587,0.05132541,0.6949145],"study_design_scores_gemma":[0.00014512915,0.0003944293,0.034604836,0.015208948,0.00057440763,0.0015909105,0.0052071353,0.027922155,0.00087124546,0.6756179,0.23759602,0.0002669416],"about_ca_topic_score_codex":0.009617447,"about_ca_topic_score_gemma":0.00423084,"teacher_disagreement_score":0.10752045,"about_ca_system_score_codex":0.00437325,"about_ca_system_score_gemma":0.006513754,"threshold_uncertainty_score":0.56862926},"labels":[],"label_agreement":null},{"id":"W7099046002","doi":"","title":"Framing [Con]text","year":2016,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Foregrounding; Graffiti; Framing (construction); Reading (process); Context (archaeology); Space (punctuation)","score_opus":0.012636405230854589,"score_gpt":0.25970565556770125,"score_spread":0.24706925033684668,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099046002","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0067486446,0.0006683013,0.01683118,0.010487592,0.0022199408,0.00011488899,0.00044354965,0.00047827384,0.9620076],"genre_scores_gemma":[0.42996368,0.0017487041,0.01589669,0.0051535587,0.0027687645,0.00027578734,0.0015712589,0.0016444108,0.5409771],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99729365,0.0011355648,0.00012527294,0.00047973348,0.00068159396,0.0002841496],"domain_scores_gemma":[0.99737525,0.00080019876,0.00028849326,0.00070398283,0.0006776325,0.00015448972],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0022518665,0.0007626964,0.00034153525,0.0023295747,0.003967524,0.009466834,0.0013263747,0.0018525119,0.06278921],"category_scores_gemma":[0.0069605582,0.00027516656,0.00036922374,0.0022670028,0.007706269,0.0090356795,0.004233975,0.0025016633,0.017256811],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002363094,0.000008915446,0.00022389003,0.00010688262,0.0000056276017,0.00026518415,0.02334777,0.000079660924,0.0006012202,0.8766842,0.061380036,0.037273154],"study_design_scores_gemma":[0.000004841018,0.000008522342,0.00022708338,0.00013547373,0.000005745716,0.00016191737,0.004912974,0.00012851138,0.0004275803,0.032445572,0.9615308,0.000010962039],"about_ca_topic_score_codex":0.0069074193,"about_ca_topic_score_gemma":0.007877049,"teacher_disagreement_score":0.9372108,"about_ca_system_score_codex":0.003710724,"about_ca_system_score_gemma":0.003018192,"threshold_uncertainty_score":0.21005082},"labels":[],"label_agreement":null},{"id":"W7099282365","doi":"","title":"Directions for research and development on electronic portfolios. Canadian Journal of Learning and","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Variety (cybernetics); Electronic portfolio; Portfolio; Key (lock); Electronic publishing; Process (computing); Electronic learning","score_opus":0.042934705225769386,"score_gpt":0.3526028361509311,"score_spread":0.3096681309251617,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099282365","genre_codex":"commentary","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064514442,0.15522553,0.0065271403,0.6233239,0.011887973,0.00027722452,0.0008135092,0.00037973933,0.19511354],"genre_scores_gemma":[0.21364373,0.5074192,0.06001926,0.047095273,0.0050086966,0.0004371105,0.0021110543,0.00033645245,0.16392918],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9888819,0.0019491836,0.000773197,0.0004359744,0.006695813,0.0012639127],"domain_scores_gemma":[0.9418938,0.0122281015,0.002082448,0.0019448553,0.03125968,0.010591036],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020025369,0.0005443728,0.00092725904,0.007789601,0.007003464,0.015546503,0.0030838079,0.004194017,0.040882424],"category_scores_gemma":[0.0363428,0.0004345895,0.0005010964,0.020296173,0.009096575,0.017829916,0.0046281954,0.0031560322,0.004193515],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004257093,0.000119311284,0.002042517,0.0016117165,0.000012369496,0.00019445452,0.0030869336,0.00031255142,0.00020215637,0.16189584,0.35909104,0.47138855],"study_design_scores_gemma":[0.000015630023,0.00002563481,0.005167324,0.0015332499,0.000009650311,0.00014462392,0.007662957,0.0002545892,0.00018725313,0.028925262,0.9560217,0.000052018797],"about_ca_topic_score_codex":0.6513232,"about_ca_topic_score_gemma":0.7366847,"teacher_disagreement_score":0.3486768,"about_ca_system_score_codex":0.035623237,"about_ca_system_score_gemma":0.17711174,"threshold_uncertainty_score":0.7014605},"labels":[],"label_agreement":null},{"id":"W7099331221","doi":"","title":"Web Site: www.cprn.org Too Many Left Behind: Canada’s Adult Education and Training System","year":2006,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Training (meteorology); Adult education; Training system; The Internet; Information system; Vocational education","score_opus":0.007423510758656562,"score_gpt":0.22650636735881394,"score_spread":0.21908285660015736,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099331221","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010755975,0.00067324965,0.0033257422,0.012288893,0.00028700163,0.00041891675,0.49724954,0.013021719,0.4619789],"genre_scores_gemma":[0.06937235,0.0014315023,0.013699343,0.004514134,0.00013866017,0.00028926667,0.31833547,0.002490196,0.58972913],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9995005,0.000023030858,0.000022235055,0.000048361726,0.00026184972,0.00014387877],"domain_scores_gemma":[0.9956463,0.0003330177,0.000102193466,0.00016374355,0.0025362535,0.0012184023],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006580298,0.00037710028,0.00032100966,0.0028965152,0.002966169,0.0021918297,0.0010967035,0.0009310902,0.20270039],"category_scores_gemma":[0.0031108784,0.00031726764,0.0002509577,0.004624998,0.0005246029,0.0012770413,0.0010972641,0.00068473595,0.06895108],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000026206508,0.000027197595,0.0028142708,0.00004397512,0.0000025473757,0.000024058514,0.00009349214,0.00007720259,0.00008143766,0.00094619463,0.9596317,0.03623173],"study_design_scores_gemma":[0.00003684369,0.000010039198,0.02061293,0.000101552854,0.000010291865,0.000050434428,0.00060017954,0.00119101,0.0004907208,0.0007326104,0.9761263,0.000037077352],"about_ca_topic_score_codex":0.9714714,"about_ca_topic_score_gemma":0.9793387,"teacher_disagreement_score":0.98511124,"about_ca_system_score_codex":0.014888756,"about_ca_system_score_gemma":0.042506084,"threshold_uncertainty_score":0.67810035},"labels":[],"label_agreement":null},{"id":"W7099350021","doi":"","title":"What is the ELD Initiative?","year":2013,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Sustainable land management; General partnership; Context (archaeology); Work (physics); Land management; Public–private partnership; Sustainable development; Politics; Food security; Order (exchange)","score_opus":0.022190850801964535,"score_gpt":0.2720217808556896,"score_spread":0.2498309300537251,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099350021","genre_codex":"commentary","genre_gemma":"commentary","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"commentary","genre_consensus":"commentary","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0035291102,0.0103926435,0.00082611636,0.88711685,0.024635518,0.000146742,0.00092813786,0.0004049281,0.07201998],"genre_scores_gemma":[0.087216064,0.017342394,0.0062535526,0.6352045,0.015168008,0.00091850775,0.0030568405,0.00054050225,0.23429973],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9864634,0.0047853976,0.0004091581,0.0009332644,0.0032977592,0.004111042],"domain_scores_gemma":[0.9445124,0.006255199,0.0018309463,0.0012648142,0.008147265,0.03798944],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.023361467,0.00057119894,0.0006751357,0.0018585923,0.005987892,0.020995326,0.0031990162,0.012251461,0.05209172],"category_scores_gemma":[0.028892253,0.00031907062,0.0006276958,0.0014955635,0.0048716064,0.016314825,0.008030316,0.0076932875,0.019516123],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0001587642,0.00023694112,0.0031887991,0.000549739,0.000012314514,0.0003309295,0.0007857192,0.000049055958,0.0002504113,0.04050152,0.8693782,0.084557585],"study_design_scores_gemma":[0.000030550695,0.00005752552,0.0017961748,0.0007128716,0.0000052953037,0.000125241,0.0024277875,0.000053603824,0.00011346708,0.004537443,0.9901091,0.000030963987],"about_ca_topic_score_codex":0.01030034,"about_ca_topic_score_gemma":0.010612313,"teacher_disagreement_score":0.05209172,"about_ca_system_score_codex":0.0074696597,"about_ca_system_score_gemma":0.037315693,"threshold_uncertainty_score":0.17426413},"labels":[],"label_agreement":null},{"id":"W7099353464","doi":"","title":"4 See Appendix for a listing of the members of the Cerebrolysin Study Group","year":2004,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Cerebrolysin; Neuroprotection; Stroke (engine); Placebo; Adverse effect; Concomitant; Clinical trial; Pentoxifylline","score_opus":0.016875754783186828,"score_gpt":0.27211016605717087,"score_spread":0.25523441127398405,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099353464","genre_codex":"dataset","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010637499,0.0028016593,0.004299029,0.0017562666,0.0011976818,0.004181487,0.9246513,0.00067878043,0.04979615],"genre_scores_gemma":[0.06845327,0.0062791435,0.0078069153,0.0058721,0.0020467152,0.041068796,0.72258866,0.00086049584,0.14502393],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993641,0.0002006031,0.00016155108,0.00009249906,0.000087236585,0.00009403585],"domain_scores_gemma":[0.9973054,0.000654851,0.0004072462,0.00022732929,0.00091460766,0.0004906313],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0012193704,0.0011541069,0.0015226731,0.0023669591,0.0006305351,0.0007941747,0.0008022395,0.0007638009,0.27063632],"category_scores_gemma":[0.0056661586,0.0004821056,0.00057152344,0.003625323,0.00010112029,0.000851199,0.00049441494,0.0009264378,0.11225329],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0038190654,0.00040444703,0.02717951,0.0013158359,0.00023571227,0.00014774529,0.00008145711,0.000231957,0.0002922591,0.00065980485,0.90813947,0.057492744],"study_design_scores_gemma":[0.009881563,0.0017242858,0.12753946,0.0012785716,0.00048077942,0.00094705704,0.00028516437,0.00071167876,0.0004283013,0.004243772,0.8523572,0.00012222989],"about_ca_topic_score_codex":0.002662738,"about_ca_topic_score_gemma":0.0035178445,"teacher_disagreement_score":0.7293637,"about_ca_system_score_codex":0.0005264382,"about_ca_system_score_gemma":0.001000036,"threshold_uncertainty_score":0.9053687},"labels":[],"label_agreement":null},{"id":"W7099355823","doi":"","title":"Athabasca University CENTRE FOR DISTANCE EDUCATION Online Software Evaluation Report TITLE: OS Software: an alternative to costly Learning Management Systems","year":2003,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Software; Distance education; Server; Learning Management; Download; Management system; Information technology; Order (exchange); Community college; Software system","score_opus":0.02251554135257924,"score_gpt":0.30310998035372,"score_spread":0.2805944390011408,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099355823","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.602528,0.0025753644,0.023835866,0.008663065,0.001152712,0.007993476,0.015646005,0.004547501,0.33305806],"genre_scores_gemma":[0.7571622,0.0013230118,0.05011818,0.0012904265,0.0001690452,0.005555366,0.015028314,0.0011654558,0.16818798],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9883923,0.0033459344,0.00060497696,0.0006524651,0.0063931923,0.0006110421],"domain_scores_gemma":[0.9569255,0.010475406,0.0011245274,0.0023385563,0.023342272,0.005793707],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0075940588,0.0004010572,0.000545267,0.0031431331,0.0013013242,0.003044623,0.0009788619,0.0006186851,0.033500284],"category_scores_gemma":[0.019564789,0.00018588958,0.0003234337,0.0026201946,0.00049380667,0.0021963029,0.0017163745,0.0009287883,0.0052524623],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022983542,0.0046672355,0.023743533,0.00069006707,0.000066768145,0.00015733407,0.0009568319,0.0009818766,0.0046402737,0.0065152743,0.13518004,0.82010245],"study_design_scores_gemma":[0.0022322424,0.012137895,0.2831691,0.00072138145,0.00025925992,0.0005967484,0.004250715,0.013826155,0.027864851,0.0037619297,0.65092695,0.000252713],"about_ca_topic_score_codex":0.021421252,"about_ca_topic_score_gemma":0.026402557,"teacher_disagreement_score":0.033500284,"about_ca_system_score_codex":0.0031676379,"about_ca_system_score_gemma":0.0046959156,"threshold_uncertainty_score":0.11206961},"labels":[],"label_agreement":null},{"id":"W7099356009","doi":"","title":"3Department of Clinical Biochemistry, University Health Network and Toronto Medical Laboratories, Toronto","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Public health; MEDLINE; Digital health; Health care","score_opus":0.03460157114186602,"score_gpt":0.3499278705393366,"score_spread":0.3153262993974706,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099356009","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16085088,0.0859861,0.035270832,0.11498047,0.011856869,0.0029721477,0.21440141,0.0076233507,0.36605805],"genre_scores_gemma":[0.40135506,0.032885484,0.050208643,0.013632943,0.0039084954,0.0016605353,0.04756092,0.0025781202,0.44620973],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9981969,0.00028660372,0.00016469885,0.00056149234,0.00044101218,0.00034930446],"domain_scores_gemma":[0.9882449,0.0022568903,0.0006575063,0.0007321406,0.004337382,0.0037713202],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.001503419,0.0014518773,0.0010568595,0.0035248902,0.0029019422,0.0028305491,0.0019480351,0.0017996979,0.1461711],"category_scores_gemma":[0.007954789,0.00081085425,0.0008896608,0.0034534298,0.0009920128,0.001140456,0.001583045,0.0016390036,0.022549003],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.003991404,0.0003792318,0.080984235,0.005221392,0.00037339405,0.009287225,0.0038250757,0.002223113,0.01488415,0.013312048,0.5777034,0.2878153],"study_design_scores_gemma":[0.000652412,0.00043754987,0.20705934,0.0015760186,0.0005309262,0.004327658,0.0023334823,0.004106712,0.011985305,0.0061535416,0.7606691,0.00016798828],"about_ca_topic_score_codex":0.23422581,"about_ca_topic_score_gemma":0.45662975,"teacher_disagreement_score":0.8538289,"about_ca_system_score_codex":0.017986918,"about_ca_system_score_gemma":0.018856397,"threshold_uncertainty_score":0.48899102},"labels":[],"label_agreement":null},{"id":"W7099358921","doi":"","title":"thesis, in preparation","year":2011,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Notice; Table (database); Work (physics); Government (linguistics); Production (economics); Schedule","score_opus":0.037197443549917726,"score_gpt":0.2782497154812454,"score_spread":0.24105227193132767,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099358921","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0029947965,0.007313088,0.005453043,0.06372079,0.10754248,0.000736323,0.013493437,0.0014554107,0.7972907],"genre_scores_gemma":[0.0058107995,0.0019521136,0.0012366968,0.002817658,0.00549451,0.00015630467,0.002276091,0.00035855244,0.9798972],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99881065,0.00014443818,0.00007491153,0.00025399588,0.00058095256,0.00013512539],"domain_scores_gemma":[0.9961914,0.00025579587,0.00011758403,0.00027986124,0.0022277734,0.0009276968],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.001581799,0.00086449494,0.0007735688,0.0012848795,0.0025035085,0.0067203245,0.0012216409,0.0013864124,0.5111656],"category_scores_gemma":[0.0066437367,0.0003743441,0.0007267704,0.0012392375,0.0008229163,0.0025966812,0.002447974,0.00267914,0.3474624],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002947659,0.0000305812,0.00013504668,0.00007149459,0.0000029116636,0.00004989504,0.00020205103,0.000058805475,0.00022564897,0.0071507893,0.9549352,0.037108168],"study_design_scores_gemma":[0.0000055957767,0.000024387933,0.0004294074,0.000073502466,0.0000021074993,0.00005195025,0.00016687474,0.000035012417,0.0000800957,0.0014666271,0.99766135,0.0000030316864],"about_ca_topic_score_codex":0.0038820473,"about_ca_topic_score_gemma":0.007200844,"teacher_disagreement_score":0.48883438,"about_ca_system_score_codex":0.0031112598,"about_ca_system_score_gemma":0.004243301,"threshold_uncertainty_score":0.6972629},"labels":[],"label_agreement":null},{"id":"W7099364877","doi":"","title":"of LaborSelf-Selection, Immigrant Public Finance Performance and Canadian Citizenship","year":2005,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Citizenship; Immigration; Public policy; Work (physics); Politics; Center (category theory)","score_opus":0.010987914678989363,"score_gpt":0.22076052164071644,"score_spread":0.2097726069617271,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099364877","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8777344,0.0017214525,0.00021388166,0.01684001,0.000099374796,0.00004478477,0.0020954309,0.000030915624,0.101219624],"genre_scores_gemma":[0.99216837,0.00034836368,0.00004725943,0.00020676825,0.000013094902,0.0000058689825,0.00032160233,0.000007888885,0.006880648],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9973273,0.00033276717,0.000063994725,0.00016810335,0.0007228962,0.0013850002],"domain_scores_gemma":[0.99122405,0.0014754207,0.0011856649,0.00041086515,0.0023525653,0.0033515869],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0027245602,0.00024655915,0.000485659,0.002415745,0.0076143285,0.0044187834,0.0012659892,0.0009101094,0.017102279],"category_scores_gemma":[0.012580926,0.00012862572,0.00043824024,0.0050258734,0.0032964386,0.0016540578,0.002571004,0.0013524842,0.00050813955],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00023964177,0.00019334682,0.8360418,0.00005272597,0.00010254059,0.00020006711,0.019360017,0.00063308165,0.00007751369,0.064683884,0.029154621,0.049260743],"study_design_scores_gemma":[0.00002819419,0.00005388661,0.90985775,0.0001743241,0.000062759485,0.00008211784,0.047006235,0.0013378244,0.0001270278,0.007303298,0.03389441,0.000072217954],"about_ca_topic_score_codex":0.9863132,"about_ca_topic_score_gemma":0.9873453,"teacher_disagreement_score":0.03850764,"about_ca_system_score_codex":0.03850764,"about_ca_system_score_gemma":0.05574271,"threshold_uncertainty_score":0.2793938},"labels":[],"label_agreement":null},{"id":"W7099366164","doi":"","title":"Treatment Assurance and Coordination of Care in Canadian and American Tuberculosis Control Systems","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Tuberculosis; Tuberculosis control; Public health; Quality assurance; Health care; Control (management); Closure (psychology)","score_opus":0.008604863539453366,"score_gpt":0.2514704501545439,"score_spread":0.24286558661509053,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099366164","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5169279,0.011670169,0.0147186415,0.1394767,0.001125008,0.0011454739,0.0054483847,0.00089024496,0.30859745],"genre_scores_gemma":[0.9660985,0.0029666582,0.011328011,0.0046046325,0.00012789742,0.00017719313,0.0015954516,0.00007172952,0.013029914],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9863479,0.0026016321,0.00063768326,0.0012202411,0.0051522884,0.004040192],"domain_scores_gemma":[0.9730452,0.0029003404,0.0017989544,0.00086597167,0.015402923,0.0059865867],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.010091573,0.00032514927,0.00038791166,0.0038580087,0.01389864,0.007962323,0.0032623734,0.0014370518,0.004324452],"category_scores_gemma":[0.02813871,0.0005100182,0.0006531407,0.008521369,0.0034134602,0.0021028926,0.0035469886,0.0018847424,0.00026543243],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00069618697,0.00046145974,0.26372865,0.0006560099,0.00028351435,0.0007218614,0.033838455,0.014402278,0.0013630951,0.2236339,0.12474582,0.33546874],"study_design_scores_gemma":[0.0001950442,0.00017797411,0.44495234,0.0008930613,0.00027243767,0.0003990168,0.030829823,0.017820574,0.0011443116,0.01742345,0.48549607,0.00039591087],"about_ca_topic_score_codex":0.99619377,"about_ca_topic_score_gemma":0.997186,"teacher_disagreement_score":0.24331145,"about_ca_system_score_codex":0.24331145,"about_ca_system_score_gemma":0.36756665,"threshold_uncertainty_score":0.87765145},"labels":[],"label_agreement":null},{"id":"W7099370898","doi":"","title":"NRC Senior Resident Inspector","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Sump (aquarium); Dominion; Plan (archaeology); Power station; Containment (computer programming); Nuclear power; Materials testing","score_opus":0.011182337174478222,"score_gpt":0.28438942390844024,"score_spread":0.27320708673396205,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099370898","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0064696716,0.0030568626,0.0073151793,0.07338811,0.047351226,0.002669599,0.018381316,0.011019057,0.83034897],"genre_scores_gemma":[0.0028703932,0.0006366627,0.0015548781,0.0051615294,0.0011039346,0.00020701095,0.0020750056,0.00059047755,0.98580015],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976301,0.0001752487,0.00012212963,0.00040965911,0.0013113654,0.0003513793],"domain_scores_gemma":[0.9878395,0.0002800214,0.00026347308,0.00053942355,0.008698937,0.0023786845],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0034680546,0.0008243031,0.0005540438,0.002385661,0.0037513447,0.0030766565,0.0020846645,0.0021154617,0.6130462],"category_scores_gemma":[0.006636685,0.00078657945,0.00045983813,0.0013505657,0.00058667036,0.0017913466,0.0027747788,0.002546512,0.32576802],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016367321,0.000025806545,0.00014123778,0.000031410113,5.9117133e-7,0.00005220234,0.000031467524,0.000011247175,0.00016104039,0.0002057902,0.97605646,0.023266379],"study_design_scores_gemma":[0.0000075688536,0.000023636176,0.00090476725,0.000034128825,0.0000016570608,0.00010907214,0.0002347973,0.000050286588,0.000121017496,0.00009383937,0.99841225,0.000006930558],"about_ca_topic_score_codex":0.03602478,"about_ca_topic_score_gemma":0.0684583,"teacher_disagreement_score":0.38695377,"about_ca_system_score_codex":0.0032497484,"about_ca_system_score_gemma":0.015702676,"threshold_uncertainty_score":0.5519426},"labels":[],"label_agreement":null},{"id":"W7099906263","doi":"","title":"PROBABILISTIC MODELING AND BAYESIAN INFERENCE OF METAL- LOSS CORROSION WITH APPLICATION IN RELIABILITY ANAYSIS FOR ENERGY PIPELINES","year":2014,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Reliability (semiconductor); Pipeline transport; Probabilistic logic; Pipeline (software); Integrity management; Bayesian probability; Energy (signal processing)","score_opus":0.010336909110076438,"score_gpt":0.2561225310993347,"score_spread":0.24578562198925827,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099906263","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.033570018,0.0005789354,0.96257997,0.0005575143,0.000041907548,0.00004117804,0.00030274925,0.00044552682,0.0018822905],"genre_scores_gemma":[0.7952406,0.0011476006,0.19672206,0.00014469076,0.00015671052,0.0001709284,0.00091204117,0.00027107354,0.0052343113],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9989236,0.00041771564,0.00006679887,0.00026581617,0.0002428962,0.00008316601],"domain_scores_gemma":[0.9908557,0.0075671086,0.0005586262,0.0003261108,0.00057533855,0.000117127245],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0040809354,0.0006860158,0.0013515693,0.0016333907,0.00074389216,0.0021813968,0.0021074633,0.0015856866,0.0028392908],"category_scores_gemma":[0.02298758,0.0012114247,0.0016463688,0.0016490167,0.0011429002,0.0026842975,0.0013236835,0.002125305,0.00057406595],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006577168,0.000046776928,0.0016114479,0.00007747566,0.0000564569,0.000055347922,0.00010232173,0.94692093,0.00059119595,0.023586037,0.0008984507,0.02598785],"study_design_scores_gemma":[0.0000044676426,0.0000073775186,0.00028835036,0.000007803888,0.000008118165,0.000013966238,0.000008835883,0.985788,0.00014730475,0.013463974,0.00025547112,0.000006369014],"about_ca_topic_score_codex":0.020838983,"about_ca_topic_score_gemma":0.020815086,"teacher_disagreement_score":0.020838983,"about_ca_system_score_codex":0.0019738644,"about_ca_system_score_gemma":0.0018633344,"threshold_uncertainty_score":0.04143536},"labels":[],"label_agreement":null},{"id":"W7099911052","doi":"","title":"Cytogenet Genome Res 109:415–479 (2005) DOI: 10.1159/000084205 Second Report on Chicken Genes and Chromosomes 2005 Organized by","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Genome; Gene; Chromosome; Inheritance (genetic algorithm); Human genome","score_opus":0.013679011497843558,"score_gpt":0.24117399172527196,"score_spread":0.2274949802274284,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099911052","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06472328,0.0789039,0.032138582,0.022522824,0.008873205,0.0008467933,0.10544035,0.018828908,0.6677222],"genre_scores_gemma":[0.12719569,0.04493801,0.06256568,0.0021785821,0.0020450784,0.0005661109,0.15159452,0.0032111926,0.60570514],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99962425,0.00003696321,0.000030169296,0.00010535091,0.00015380024,0.000049413484],"domain_scores_gemma":[0.99889994,0.00030275143,0.00009020366,0.00025736485,0.00025131716,0.00019828016],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.0013020321,0.0007848159,0.0012515241,0.0026615544,0.00091523957,0.0040303203,0.00075679406,0.0013377902,0.31203052],"category_scores_gemma":[0.0023938504,0.00074285927,0.0006036861,0.0032805046,0.0014715805,0.0024395764,0.0018397148,0.0020571572,0.2990423],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0004355046,0.00025293423,0.00483697,0.00048594267,0.00007751964,0.00068597397,0.0002270466,0.00033066794,0.04967736,0.009079795,0.41791096,0.5159994],"study_design_scores_gemma":[0.00014443893,0.000086917,0.02412856,0.00020434315,0.00006876042,0.0014235288,0.00013485123,0.0007084221,0.014328299,0.008063936,0.9506666,0.000041276857],"about_ca_topic_score_codex":0.0024163497,"about_ca_topic_score_gemma":0.0033289888,"teacher_disagreement_score":0.68796945,"about_ca_system_score_codex":0.00074158364,"about_ca_system_score_gemma":0.00078481436,"threshold_uncertainty_score":0.98130494},"labels":[],"label_agreement":null},{"id":"W7099913945","doi":"","title":"Highlights","year":2007,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Remand (court procedure); Community service; Suicide prevention; Injury prevention; Poison control; Occupational safety and health","score_opus":0.010939475431344952,"score_gpt":0.2730963889082409,"score_spread":0.26215691347689596,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099913945","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0022586102,0.016984833,0.0023927838,0.1337626,0.07179611,0.0007542084,0.01732785,0.0020022078,0.75272083],"genre_scores_gemma":[0.0214096,0.025513172,0.0030410562,0.062222634,0.028328577,0.0007325134,0.02320757,0.0011502225,0.8343947],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9980014,0.00034256323,0.00012859101,0.0003256831,0.00077220093,0.0004295079],"domain_scores_gemma":[0.9956428,0.0003011134,0.00020113071,0.00024241011,0.0019036671,0.0017088105],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0017392634,0.0009023348,0.00066702167,0.0014136557,0.0024177823,0.0054450715,0.0025731728,0.0039712493,0.50967413],"category_scores_gemma":[0.008437907,0.0003356108,0.00095107115,0.0012202396,0.00071637065,0.0028201444,0.003977915,0.0036958507,0.25829825],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003297605,0.00002502103,0.0003965099,0.00023848508,0.0000057776692,0.00021065833,0.00006483608,0.000022238297,0.000071609684,0.0023631887,0.9425521,0.05401659],"study_design_scores_gemma":[0.0000102984495,0.000014407773,0.00060271117,0.00029976034,0.0000038039568,0.00031355527,0.000120200595,0.00001619097,0.000046447334,0.0009123095,0.9976554,0.0000048399556],"about_ca_topic_score_codex":0.00844399,"about_ca_topic_score_gemma":0.007308539,"teacher_disagreement_score":0.49032587,"about_ca_system_score_codex":0.0030124234,"about_ca_system_score_gemma":0.011063569,"threshold_uncertainty_score":0},"labels":[],"label_agreement":null},{"id":"W7099915710","doi":"","title":"Urogynaecology of the Society of Obstetricians and","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Urinary incontinence; Sling (weapon); Task force; Stress incontinence; MEDLINE","score_opus":0.02633763762919184,"score_gpt":0.25475111272079676,"score_spread":0.22841347509160492,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7099915710","genre_codex":"review","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00888695,0.63817793,0.002404147,0.15222298,0.054045577,0.0003143022,0.0049025286,0.0003363248,0.13870929],"genre_scores_gemma":[0.08687529,0.7448919,0.012713824,0.036439985,0.029178476,0.00074975414,0.007218916,0.00024723983,0.08168463],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9968765,0.0009159824,0.00071356207,0.00026737977,0.0008391494,0.00038742658],"domain_scores_gemma":[0.9851681,0.0028118629,0.0017771462,0.0006066553,0.007015366,0.0026209098],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00380201,0.0005529723,0.000873485,0.0046254704,0.0011783548,0.0018491643,0.0009319884,0.0014695737,0.03804083],"category_scores_gemma":[0.016595146,0.0003775096,0.00070560595,0.0042791124,0.00082301593,0.0012031541,0.0021529552,0.0018626072,0.008387474],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008661267,0.000044545835,0.013947004,0.0034285276,0.000074960015,0.0007515013,0.0005603129,0.000070199894,0.00038730283,0.0024262995,0.4359662,0.54225653],"study_design_scores_gemma":[0.00002793715,0.00003976445,0.020084934,0.0048478814,0.000048326805,0.001932006,0.0005049751,0.000025885935,0.000092395174,0.0006851003,0.9716939,0.000016965647],"about_ca_topic_score_codex":0.008090037,"about_ca_topic_score_gemma":0.0152023565,"teacher_disagreement_score":0.9619592,"about_ca_system_score_codex":0.0018986941,"about_ca_system_score_gemma":0.012931533,"threshold_uncertainty_score":0.12725931},"labels":[],"label_agreement":null},{"id":"W7100354937","doi":"","title":"Data and Text Mining METIS: multiple extraction techniques for informative sentences","year":2008,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Annotation; Support vector machine; Sentence; Component (thermodynamics); Information extraction; Text mining; Metis","score_opus":0.06496683998000835,"score_gpt":0.3358570197140637,"score_spread":0.2708901797340553,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7100354937","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.024405254,0.0012576993,0.8927314,0.0014777279,0.0003022452,0.0018877082,0.04371061,0.031007817,0.0032194187],"genre_scores_gemma":[0.048134677,0.00033955422,0.9064739,0.00015853404,0.00015914724,0.0011505149,0.04101146,0.0008870419,0.0016851993],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9955089,0.0010912061,0.0010425735,0.00094917894,0.0012588437,0.00014929527],"domain_scores_gemma":[0.98469687,0.00802515,0.0016230701,0.001777718,0.0035483548,0.00032881435],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0057249977,0.001616409,0.0011970475,0.01031194,0.0013672002,0.002256404,0.0018117677,0.0011297405,0.007401329],"category_scores_gemma":[0.018787717,0.0007816044,0.001560485,0.006567107,0.0007331696,0.0028227537,0.002458669,0.0020389762,0.0057096877],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00066288566,0.00031563474,0.006520664,0.003651277,0.00037449074,0.0012796838,0.001968012,0.0025865724,0.05322624,0.01011249,0.08004003,0.8392621],"study_design_scores_gemma":[0.0005211541,0.000974093,0.027404888,0.0012906715,0.0013660091,0.006311987,0.0033391046,0.3510339,0.2366614,0.060730733,0.30985886,0.00050724624],"about_ca_topic_score_codex":0.001204018,"about_ca_topic_score_gemma":0.002716136,"teacher_disagreement_score":0.01031194,"about_ca_system_score_codex":0.0007798605,"about_ca_system_score_gemma":0.0022635325,"threshold_uncertainty_score":0.030277014},"labels":[],"label_agreement":null},{"id":"W7100711346","doi":"","title":"Background The Organization and Its Products","year":2015,"lang":"en","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Software; Corporation; Product (mathematics); The Internet; Pharmaceutical industry; New product development; Software development; Observational study","score_opus":0.04362676574260613,"score_gpt":0.27215468775286455,"score_spread":0.22852792201025843,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7100711346","genre_codex":"other","genre_gemma":"other","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"other","genre_consensus":"other","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010034509,0.0068154857,0.023754274,0.03506691,0.002245984,0.0021185628,0.021802157,0.0034533797,0.8947087],"genre_scores_gemma":[0.1127447,0.012300077,0.07691725,0.010935294,0.0017722908,0.0029746168,0.050806507,0.0040131104,0.7275362],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99027914,0.0015903185,0.0005368162,0.0013222747,0.0048136623,0.0014578496],"domain_scores_gemma":[0.9617714,0.0038099803,0.0017196517,0.0034002364,0.023420876,0.0058777523],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.008527764,0.0006622236,0.00040956726,0.002318275,0.0033393505,0.010735674,0.0014645006,0.0018446827,0.075759545],"category_scores_gemma":[0.026551964,0.00054625905,0.00029456313,0.004756272,0.0015514526,0.0036008079,0.0021459667,0.0017645295,0.051655043],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00017195786,0.00022650564,0.0037247918,0.0004622063,0.0000070769884,0.0002057132,0.00096680236,0.00044778176,0.0011444299,0.08894154,0.664964,0.23873717],"study_design_scores_gemma":[0.000007735763,0.000018650224,0.0012043344,0.000120210476,0.0000020255816,0.000052797288,0.00015760765,0.0000608956,0.00014048857,0.0010177055,0.9972095,0.000008028477],"about_ca_topic_score_codex":0.10113976,"about_ca_topic_score_gemma":0.054042984,"teacher_disagreement_score":0.92424047,"about_ca_system_score_codex":0.012611944,"about_ca_system_score_gemma":0.034987357,"threshold_uncertainty_score":0.25344092},"labels":[],"label_agreement":null},{"id":"W7104443750","doi":"10.5281/zenodo.17559036","title":"An LLM-based Toolbox for Automated Text Mining on the Uses of Chemicals","year":2025,"lang":"","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"European Commission","keywords":"Toolbox; Software; Automation; Expert system; Visualization","score_opus":0.03712296480879256,"score_gpt":0.3008059306487143,"score_spread":0.26368296583992173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7104443750","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050553186,0.00036806875,0.6655753,0.000517533,0.00012302272,0.00047548715,0.049708597,0.27462274,0.0035539719],"genre_scores_gemma":[0.036426034,0.0004481132,0.8687901,0.0003640054,0.00006295849,0.000848206,0.08238331,0.0058194366,0.0048577674],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987117,0.00024861947,0.00027592352,0.0003621178,0.00034031903,0.000061285566],"domain_scores_gemma":[0.99654657,0.0018923099,0.00030987413,0.0005376495,0.0005091422,0.00020447765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019049384,0.0019593262,0.0009356478,0.0055427263,0.000663476,0.002675589,0.001416985,0.0010695873,0.026070092],"category_scores_gemma":[0.0072137523,0.00070520194,0.0018543175,0.0027344534,0.00039462285,0.0029434087,0.0027025884,0.0015332227,0.019334026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067486457,0.0004064382,0.0052972217,0.0034191553,0.00037713756,0.0007672123,0.0007541133,0.0076710917,0.031941753,0.012593847,0.16691034,0.76918685],"study_design_scores_gemma":[0.00032033087,0.0002653851,0.008192095,0.0009492004,0.00028499757,0.0014880276,0.0007321574,0.44114047,0.06851006,0.05830088,0.4196309,0.00018555955],"about_ca_topic_score_codex":0.0022705637,"about_ca_topic_score_gemma":0.004355871,"teacher_disagreement_score":0.026070092,"about_ca_system_score_codex":0.0008615875,"about_ca_system_score_gemma":0.0016164652,"threshold_uncertainty_score":0.08721316},"labels":[],"label_agreement":null},{"id":"W7104467956","doi":"10.5281/zenodo.17559035","title":"An LLM-based Toolbox for Automated Text Mining on the Uses of Chemicals","year":2025,"lang":"","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"European Commission","keywords":"Toolbox; Software; Automation; Expert system; Visualization","score_opus":0.03712296480879256,"score_gpt":0.3008059306487143,"score_spread":0.26368296583992173,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7104467956","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0050553186,0.00036806875,0.6655753,0.000517533,0.00012302272,0.00047548715,0.049708597,0.27462274,0.0035539719],"genre_scores_gemma":[0.036426034,0.0004481132,0.8687901,0.0003640054,0.00006295849,0.000848206,0.08238331,0.0058194366,0.0048577674],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9987117,0.00024861947,0.00027592352,0.0003621178,0.00034031903,0.000061285566],"domain_scores_gemma":[0.99654657,0.0018923099,0.00030987413,0.0005376495,0.0005091422,0.00020447765],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019049384,0.0019593262,0.0009356478,0.0055427263,0.000663476,0.002675589,0.001416985,0.0010695873,0.026070092],"category_scores_gemma":[0.0072137523,0.00070520194,0.0018543175,0.0027344534,0.00039462285,0.0029434087,0.0027025884,0.0015332227,0.019334026],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00067486457,0.0004064382,0.0052972217,0.0034191553,0.00037713756,0.0007672123,0.0007541133,0.0076710917,0.031941753,0.012593847,0.16691034,0.76918685],"study_design_scores_gemma":[0.00032033087,0.0002653851,0.008192095,0.0009492004,0.00028499757,0.0014880276,0.0007321574,0.44114047,0.06851006,0.05830088,0.4196309,0.00018555955],"about_ca_topic_score_codex":0.0022705637,"about_ca_topic_score_gemma":0.004355871,"teacher_disagreement_score":0.026070092,"about_ca_system_score_codex":0.0008615875,"about_ca_system_score_gemma":0.0016164652,"threshold_uncertainty_score":0.08721316},"labels":[],"label_agreement":null},{"id":"W7108210586","doi":"10.5281/zenodo.17783407","title":"Building Shared Language for Salmon Knowledge","year":2025,"lang":"enc","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Fisheries and Oceans Canada","funders":"","keywords":"Blueprint; Terminology; Presentation (obstetrics); Ontology; Knowledge sharing; Alliance","score_opus":0.02776173884229347,"score_gpt":0.30813211794058354,"score_spread":0.28037037909829005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7108210586","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002713917,0.00023915243,0.975743,0.0033983707,0.00032374763,0.00050875783,0.0026693114,0.0080662435,0.006337544],"genre_scores_gemma":[0.028296284,0.0004559095,0.9386771,0.0013181426,0.00013907769,0.0012438456,0.019451521,0.003292151,0.00712592],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.97767705,0.007938465,0.004743218,0.0037395053,0.004822422,0.0010793733],"domain_scores_gemma":[0.9628373,0.011880861,0.0016031744,0.012950305,0.008746631,0.001981764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029330112,0.0013857444,0.0020624422,0.006248465,0.0051364666,0.012400392,0.0066293143,0.004269974,0.012440048],"category_scores_gemma":[0.04123491,0.0025707055,0.0055418885,0.0067507587,0.0056392793,0.03129328,0.035850182,0.008210819,0.008240087],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022617297,0.00033128585,0.0015793388,0.0011327689,0.00020299591,0.00093869684,0.007929847,0.007573383,0.0057140025,0.7145522,0.084094316,0.17572504],"study_design_scores_gemma":[0.000101253965,0.00009124654,0.00036342838,0.00079991575,0.00014854698,0.000356097,0.002423883,0.04560651,0.007896358,0.4311719,0.5108681,0.00017279078],"about_ca_topic_score_codex":0.012099684,"about_ca_topic_score_gemma":0.013670703,"teacher_disagreement_score":0.029330112,"about_ca_system_score_codex":0.004073512,"about_ca_system_score_gemma":0.014008763,"threshold_uncertainty_score":0.15511435},"labels":[],"label_agreement":null},{"id":"W7108217307","doi":"10.5281/zenodo.17783406","title":"Building Shared Language for Salmon Knowledge","year":2025,"lang":"enc","type":"article","venue":"Zenodo (CERN European Organization for Nuclear Research)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Fisheries and Oceans Canada","funders":"","keywords":"Blueprint; Terminology; Presentation (obstetrics); Ontology; Knowledge sharing; Alliance","score_opus":0.02776173884229347,"score_gpt":0.30813211794058354,"score_spread":0.28037037909829005,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7108217307","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.002713917,0.00023915243,0.975743,0.0033983707,0.00032374763,0.00050875783,0.0026693114,0.0080662435,0.006337544],"genre_scores_gemma":[0.028296284,0.0004559095,0.9386771,0.0013181426,0.00013907769,0.0012438456,0.019451521,0.003292151,0.00712592],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.97767705,0.007938465,0.004743218,0.0037395053,0.004822422,0.0010793733],"domain_scores_gemma":[0.9628373,0.011880861,0.0016031744,0.012950305,0.008746631,0.001981764],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.029330112,0.0013857444,0.0020624422,0.006248465,0.0051364666,0.012400392,0.0066293143,0.004269974,0.012440048],"category_scores_gemma":[0.04123491,0.0025707055,0.0055418885,0.0067507587,0.0056392793,0.03129328,0.035850182,0.008210819,0.008240087],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00022617297,0.00033128585,0.0015793388,0.0011327689,0.00020299591,0.00093869684,0.007929847,0.007573383,0.0057140025,0.7145522,0.084094316,0.17572504],"study_design_scores_gemma":[0.000101253965,0.00009124654,0.00036342838,0.00079991575,0.00014854698,0.000356097,0.002423883,0.04560651,0.007896358,0.4311719,0.5108681,0.00017279078],"about_ca_topic_score_codex":0.012099684,"about_ca_topic_score_gemma":0.013670703,"teacher_disagreement_score":0.029330112,"about_ca_system_score_codex":0.004073512,"about_ca_system_score_gemma":0.014008763,"threshold_uncertainty_score":0.15511435},"labels":[],"label_agreement":null},{"id":"W7108685807","doi":"10.5376/cmb.2025.15.0016","title":"Large Language Models for Biological Knowledge Extraction","year":2025,"lang":"","type":"article","venue":"Computational Molecular Biology","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Process (computing); Information extraction; Knowledge representation and reasoning; Knowledge extraction; Relation (database); Domain (mathematical analysis); Relationship extraction; Event (particle physics); Domain knowledge","score_opus":0.024366617849239807,"score_gpt":0.37268388600902347,"score_spread":0.34831726815978364,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7108685807","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.00184695,0.0013978309,0.9876651,0.0010051169,0.00009538102,0.00012603475,0.0022005432,0.003728237,0.0019348562],"genre_scores_gemma":[0.13408695,0.0033314691,0.8413804,0.0008909086,0.0002786674,0.0010330649,0.01293294,0.0008776066,0.0051880246],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99713767,0.0012676563,0.00030819097,0.0004799259,0.0006886748,0.00011780903],"domain_scores_gemma":[0.9908591,0.006703653,0.0004916481,0.0010246278,0.0008033312,0.00011773329],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0036355858,0.0013299127,0.0011860434,0.0036849754,0.00096955936,0.0033429817,0.00213508,0.0014710543,0.005251538],"category_scores_gemma":[0.017638545,0.00090843474,0.0034394243,0.004021143,0.0008231382,0.004739771,0.002812637,0.0029311783,0.0040828856],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0002375351,0.00018686567,0.0018422336,0.0009440941,0.0004915917,0.00062483706,0.00053466455,0.32774338,0.0032215538,0.19851075,0.03762245,0.42804],"study_design_scores_gemma":[0.000022223205,0.000019877067,0.00020839281,0.000082496495,0.00005790349,0.00012648864,0.00006440434,0.78361815,0.0012305424,0.19336317,0.021175446,0.000031025207],"about_ca_topic_score_codex":0.008677028,"about_ca_topic_score_gemma":0.013696335,"teacher_disagreement_score":0.008677028,"about_ca_system_score_codex":0.002076785,"about_ca_system_score_gemma":0.0027837285,"threshold_uncertainty_score":0.019227028},"labels":[],"label_agreement":null},{"id":"W7111070802","doi":"10.1371/journal.pone.0326139.s001","title":"STROBE-checklist-v6-combined-PlosMedicine.","year":2025,"lang":"","type":"article","venue":"Figshare","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Clinical trial; Health care; Alternative medicine; MEDLINE; Clinical research; Health professionals; Randomized controlled trial","score_opus":0.01911733791462814,"score_gpt":0.2919909798293353,"score_spread":0.27287364191470714,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7111070802","genre_codex":"protocol","genre_gemma":"methods","domain_codex":null,"domain_gemma":"reporting","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":"reporting","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0030541702,0.009727397,0.035814643,0.044933315,0.020196857,0.47459316,0.3644357,0.012854913,0.034389872],"genre_scores_gemma":[0.012828989,0.003153748,0.062863916,0.008353505,0.00088017684,0.8717626,0.031593382,0.002917227,0.005646395],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.7508903,0.123663984,0.09267357,0.007484604,0.021557897,0.0037296736],"domain_scores_gemma":[0.37946904,0.46145546,0.052904103,0.034840345,0.0668088,0.0045222244],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.18750791,0.004348945,0.009336253,0.017185904,0.004300893,0.01366792,0.009849033,0.009609552,0.28561112],"category_scores_gemma":[0.48718283,0.005105831,0.011951472,0.014120024,0.0073348545,0.0105692325,0.011868746,0.008901834,0.051898245],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0022910303,0.00020719407,0.0013624952,0.30209446,0.0014013337,0.0001772781,0.0010973953,0.0007118328,0.00028620963,0.013433656,0.6416891,0.035247993],"study_design_scores_gemma":[0.012187232,0.00040213167,0.007048238,0.19207111,0.0012101711,0.00069473474,0.0018027575,0.003085295,0.0011247835,0.037710726,0.7418739,0.00078893354],"about_ca_topic_score_codex":0.0033680995,"about_ca_topic_score_gemma":0.0055498565,"teacher_disagreement_score":0.8124921,"about_ca_system_score_codex":0.012282492,"about_ca_system_score_gemma":0.027507251,"threshold_uncertainty_score":0.9916485},"labels":[],"label_agreement":null},{"id":"W7111941098","doi":"","title":"Incorporating Zoning Information into Argument Mining from Biomedical Literature","year":2022,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Zoning; Argument (complex analysis); Argumentative; Identification (biology); Task (project management); Component (thermodynamics); Convention","score_opus":0.032579041316990506,"score_gpt":0.27460675452413824,"score_spread":0.24202771320714772,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7111941098","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15879133,0.0050826785,0.8076233,0.0030463934,0.0003071426,0.0006352493,0.00467113,0.012977463,0.0068653305],"genre_scores_gemma":[0.587745,0.0012633,0.39815477,0.00029097355,0.00021358299,0.00030887488,0.009405349,0.00040439307,0.0022137559],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.997988,0.0007631079,0.00024139849,0.00050623115,0.00040114217,0.00010010094],"domain_scores_gemma":[0.9873029,0.009527347,0.0010746458,0.00076622673,0.001127862,0.00020105146],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0044982624,0.0009921876,0.0008741366,0.009065325,0.00092957274,0.0018240181,0.0011933089,0.0015508011,0.0035622402],"category_scores_gemma":[0.020533413,0.00051188667,0.0016170056,0.0037403635,0.00086209335,0.0064130723,0.002636195,0.0018064316,0.0020278147],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00091468496,0.00038231476,0.015381888,0.0027370127,0.00026645977,0.0007178554,0.0027948658,0.024716323,0.039287724,0.020748174,0.014489502,0.8775632],"study_design_scores_gemma":[0.00016925618,0.00039863042,0.015897464,0.00057635194,0.00052323204,0.00080222747,0.0015860762,0.81645435,0.034880076,0.092078835,0.036495756,0.00013774834],"about_ca_topic_score_codex":0.0023566408,"about_ca_topic_score_gemma":0.004711806,"teacher_disagreement_score":0.009065325,"about_ca_system_score_codex":0.001216391,"about_ca_system_score_gemma":0.002097527,"threshold_uncertainty_score":0.023789346},"labels":[],"label_agreement":null},{"id":"W7112131225","doi":"","title":"DiscHPO: Generative Models and Sentence Transformers for the Recognition and Normalisation of Continuous and Discontinuous Phenotype Mentions","year":2025,"lang":"en","type":"article","venue":"Research Explorer (The University of Manchester)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Open Text (Canada)","funders":"","keywords":"Disjoint sets; Sentence; Transformer; Named-entity recognition; Pipeline (software); Phrase; Probabilistic logic; Knowledge base","score_opus":0.09671815083821429,"score_gpt":0.30960700692698395,"score_spread":0.21288885608876967,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7112131225","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.0128968535,0.00051491125,0.9456928,0.00068544,0.0002344414,0.00024908813,0.0034126996,0.034236968,0.0020767928],"genre_scores_gemma":[0.3401145,0.0006187181,0.62775564,0.001035023,0.00019342022,0.0005211719,0.01717019,0.0026515722,0.009939732],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9991339,0.00028033103,0.00006185572,0.00034006004,0.00013107552,0.000052879805],"domain_scores_gemma":[0.99720746,0.002013371,0.0001288974,0.00031304048,0.0002563695,0.000080811384],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0021936973,0.0016034463,0.00057391793,0.00118934,0.00040209963,0.0014224288,0.0021692305,0.0016623694,0.012474083],"category_scores_gemma":[0.007220557,0.00084943057,0.0018845075,0.00053320243,0.0007969125,0.0028193374,0.0022330626,0.0028657683,0.006518165],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0012318653,0.000310345,0.0049334886,0.0009183077,0.00038030904,0.0013168717,0.0011568684,0.14455317,0.0349312,0.026951842,0.050404172,0.7329115],"study_design_scores_gemma":[0.000038563427,0.000090816604,0.0007012231,0.000044040647,0.00004692665,0.00027820058,0.00008235389,0.9604807,0.010005707,0.01999589,0.008198567,0.000037073212],"about_ca_topic_score_codex":0.0055316463,"about_ca_topic_score_gemma":0.009405803,"teacher_disagreement_score":0.012474083,"about_ca_system_score_codex":0.0011222722,"about_ca_system_score_gemma":0.0009899597,"threshold_uncertainty_score":0.041729987},"labels":[],"label_agreement":null},{"id":"W7116798554","doi":"10.2139/ssrn.5944137","title":"Retrieval at Risk: Systematic Failures in Embedding-Based Retrieval for Nuclear Regulatory Compliance","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Reliability (semiconductor); False positive paradox; Domain (mathematical analysis); Similarity (geometry); Proxy (statistics); Scram; Precision and recall","score_opus":0.015650007669462635,"score_gpt":0.2903834509394865,"score_spread":0.27473344327002386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7116798554","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23018599,0.0019596154,0.7258421,0.005710355,0.0003759707,0.00062678667,0.0026538563,0.022965519,0.00967978],"genre_scores_gemma":[0.7823106,0.00064433145,0.20167409,0.0009421637,0.00019098759,0.00027261837,0.004249521,0.0038411014,0.005874538],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9692314,0.011487086,0.0036280444,0.0029035513,0.011256412,0.0014935713],"domain_scores_gemma":[0.79544014,0.10435672,0.009263822,0.06954527,0.02005212,0.0013419518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020485967,0.0010420678,0.0019185169,0.0042070295,0.0020181867,0.004335259,0.0029428576,0.004003205,0.005560396],"category_scores_gemma":[0.18115589,0.0010217908,0.0013883044,0.005044134,0.0036412543,0.017058553,0.007515107,0.0034978879,0.00482795],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015136827,0.0005813584,0.02157503,0.001517272,0.00033766855,0.0013702314,0.004836904,0.04285127,0.026908746,0.0747231,0.04448945,0.7792954],"study_design_scores_gemma":[0.00015765207,0.00045753116,0.0048704473,0.0004154997,0.000511663,0.0021202096,0.0029719283,0.62221247,0.08783189,0.24838896,0.029810661,0.0002511127],"about_ca_topic_score_codex":0.003822624,"about_ca_topic_score_gemma":0.0027641815,"teacher_disagreement_score":0.020485967,"about_ca_system_score_codex":0.0014283505,"about_ca_system_score_gemma":0.0046592345,"threshold_uncertainty_score":0.108341455},"labels":[],"label_agreement":null},{"id":"W7116973535","doi":"10.64898/2025.12.20.695723","title":"Application of Large Language Models for Annotating Genes into Reactome Pathways","year":2025,"lang":"en","type":"article","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto; Ontario Institute for Cancer Research","funders":"National Institutes of Health","keywords":"Workflow; Data curation; Similarity (geometry); Resource (disambiguation); Semantic similarity; Semantics (computer science)","score_opus":0.009973064382286952,"score_gpt":0.2524537599215689,"score_spread":0.24248069553928192,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7116973535","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.012336294,0.00018926828,0.96264565,0.00035660915,0.00010582013,0.0004566972,0.0049463673,0.017180407,0.0017828001],"genre_scores_gemma":[0.076015875,0.00035387126,0.9097632,0.00021232275,0.000036824487,0.00083792163,0.009956935,0.0014850575,0.0013380822],"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99699926,0.0010605103,0.0004895603,0.0007405511,0.0006083509,0.0001017605],"domain_scores_gemma":[0.99353206,0.003605467,0.0005603846,0.0011468077,0.0010252284,0.00013011297],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050214427,0.0017854755,0.000768155,0.0030901823,0.001231315,0.0029142546,0.0020323822,0.000744524,0.002856423],"category_scores_gemma":[0.012546751,0.00087808265,0.0045518875,0.0025869233,0.00070192624,0.0023135724,0.0025615431,0.0017154767,0.0015371975],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0011938141,0.000522762,0.01420436,0.003362008,0.0009277805,0.0027168426,0.0025929858,0.4859879,0.10295989,0.12762186,0.021807544,0.23610228],"study_design_scores_gemma":[0.000073241,0.00007897368,0.0008917714,0.0001291465,0.00016996787,0.00019545147,0.00021211425,0.8635103,0.043923967,0.053062286,0.037649542,0.000103173734],"about_ca_topic_score_codex":0.0069567654,"about_ca_topic_score_gemma":0.008508965,"teacher_disagreement_score":0.0069567654,"about_ca_system_score_codex":0.0026361723,"about_ca_system_score_gemma":0.0034679228,"threshold_uncertainty_score":0.026556194},"labels":[],"label_agreement":null},{"id":"W7117116563","doi":"10.2139/ssrn.5960744","title":"Retrieval at Risk: Systematic Failures in Embedding-Based Retrieval for Nuclear Regulatory Compliance","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McMaster University","funders":"","keywords":"Reliability (semiconductor); False positive paradox; Domain (mathematical analysis); Similarity (geometry); Proxy (statistics); Scram; Precision and recall","score_opus":0.015650007669462635,"score_gpt":0.2903834509394865,"score_spread":0.27473344327002386,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117116563","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.23018599,0.0019596154,0.7258421,0.005710355,0.0003759707,0.00062678667,0.0026538563,0.022965519,0.00967978],"genre_scores_gemma":[0.7823106,0.00064433145,0.20167409,0.0009421637,0.00019098759,0.00027261837,0.004249521,0.0038411014,0.005874538],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9692314,0.011487086,0.0036280444,0.0029035513,0.011256412,0.0014935713],"domain_scores_gemma":[0.79544014,0.10435672,0.009263822,0.06954527,0.02005212,0.0013419518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.020485967,0.0010420678,0.0019185169,0.0042070295,0.0020181867,0.004335259,0.0029428576,0.004003205,0.005560396],"category_scores_gemma":[0.18115589,0.0010217908,0.0013883044,0.005044134,0.0036412543,0.017058553,0.007515107,0.0034978879,0.00482795],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0015136827,0.0005813584,0.02157503,0.001517272,0.00033766855,0.0013702314,0.004836904,0.04285127,0.026908746,0.0747231,0.04448945,0.7792954],"study_design_scores_gemma":[0.00015765207,0.00045753116,0.0048704473,0.0004154997,0.000511663,0.0021202096,0.0029719283,0.62221247,0.08783189,0.24838896,0.029810661,0.0002511127],"about_ca_topic_score_codex":0.003822624,"about_ca_topic_score_gemma":0.0027641815,"teacher_disagreement_score":0.020485967,"about_ca_system_score_codex":0.0014283505,"about_ca_system_score_gemma":0.0046592345,"threshold_uncertainty_score":0.108341455},"labels":[],"label_agreement":null},{"id":"W7117468468","doi":"10.3390/bioengineering13010030","title":"Accurate Clinical Entity Recognition and Code Mapping of Anatomopathological Reports Using BioClinicalBERT Enhanced by Retrieval-Augmented Generation: A Hybrid Deep Learning Approach","year":2025,"lang":"en","type":"article","venue":"Bioengineering","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"York University","funders":"","keywords":"Deep learning; Pipeline (software); Coding (social sciences); Code (set theory); Pattern recognition (psychology)","score_opus":0.06099622682867463,"score_gpt":0.32211940332081584,"score_spread":0.26112317649214123,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7117468468","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15930192,0.0039257994,0.7334996,0.001490338,0.0004987496,0.00086661713,0.018568207,0.07474169,0.0071071037],"genre_scores_gemma":[0.37321445,0.0010129737,0.5599388,0.00072532904,0.0001971141,0.0005002179,0.055376064,0.00094435783,0.008090732],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99878234,0.00020405218,0.00012331762,0.0004778664,0.00028370853,0.0001287049],"domain_scores_gemma":[0.9984712,0.00038539778,0.00018819267,0.00038710603,0.00050663133,0.00006149805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019072156,0.0014023762,0.00075921504,0.0048853266,0.000534338,0.001477542,0.0020268615,0.0013247922,0.0031018385],"category_scores_gemma":[0.0038833586,0.00038066815,0.0014337663,0.0020678067,0.00051772036,0.0016095433,0.0019334946,0.0011372198,0.0027069654],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00032731754,0.00030373718,0.011058036,0.000501348,0.00022991658,0.0005181759,0.00032291852,0.023679882,0.025637982,0.0018249936,0.039775554,0.8958202],"study_design_scores_gemma":[0.00010421832,0.00026153424,0.016214682,0.00024799566,0.00037750325,0.0012386398,0.00039123956,0.87909496,0.05639872,0.008855182,0.036691036,0.00012439645],"about_ca_topic_score_codex":0.013834733,"about_ca_topic_score_gemma":0.019105574,"teacher_disagreement_score":0.013834733,"about_ca_system_score_codex":0.0014941578,"about_ca_system_score_gemma":0.0020446428,"threshold_uncertainty_score":0.027508438},"labels":[],"label_agreement":null},{"id":"W7133449565","doi":"","title":"Managing alignment in open source BM; to which extent does Hirschman help understanding community reactions","year":2022,"lang":"en","type":"article","venue":"ORBi UMONS","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Université de Moncton","funders":"","keywords":"Set (abstract data type); Context (archaeology); Open source; Variety (cybernetics); Work (physics); Government (linguistics)","score_opus":0.055417352665595505,"score_gpt":0.3112109381665125,"score_spread":0.255793585500917,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7133449565","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16293792,0.0012534836,0.72494656,0.021005863,0.0010382502,0.0010578538,0.0021951643,0.01991768,0.06564725],"genre_scores_gemma":[0.63323116,0.00069549313,0.33091736,0.0011281095,0.0005217064,0.00040720086,0.0037312529,0.00463906,0.024728537],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.97655356,0.009812169,0.001428907,0.0027827935,0.007893199,0.0015293488],"domain_scores_gemma":[0.9169707,0.028525077,0.010964945,0.02049541,0.020045321,0.0029986554],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.024547955,0.00070275314,0.0007940699,0.0050247866,0.003864656,0.0075331638,0.002859567,0.0021712573,0.012792806],"category_scores_gemma":[0.12968177,0.0007434237,0.0008699487,0.006360291,0.0020978125,0.023471456,0.005813091,0.0020141015,0.0051686945],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007073874,0.00048005854,0.03749616,0.0011014575,0.00029950056,0.0009406286,0.016413746,0.009926661,0.025917472,0.24823825,0.068601295,0.5898774],"study_design_scores_gemma":[0.00014606425,0.00030915317,0.019258972,0.0007564944,0.00030340496,0.0006506481,0.015556844,0.14888221,0.056955867,0.50379974,0.25298986,0.00039068886],"about_ca_topic_score_codex":0.007204933,"about_ca_topic_score_gemma":0.010413197,"teacher_disagreement_score":0.024547955,"about_ca_system_score_codex":0.0024604884,"about_ca_system_score_gemma":0.006333572,"threshold_uncertainty_score":0.1298235},"labels":[],"label_agreement":null},{"id":"W7134171330","doi":"10.1109/bigdata66926.2025.11401680","title":"From Text to Insight: Towards Robust RAG Pipelines for Transcript-Based Clinical Screening","year":2025,"lang":"","type":"article","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Waterloo","funders":"","keywords":"Pipeline transport; Pipeline (software); Noise (video); Key (lock)","score_opus":0.0747946657671624,"score_gpt":0.37328498513361486,"score_spread":0.29849031936645243,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7134171330","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.015291295,0.0021419004,0.7966753,0.0024955266,0.00035420127,0.0007026877,0.053761642,0.12438424,0.004193196],"genre_scores_gemma":[0.12107328,0.0015288275,0.76903284,0.001686353,0.00030407292,0.0006827669,0.09651811,0.004693141,0.0044805673],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99665534,0.0007983821,0.00036504486,0.0010709699,0.0008195467,0.00029077422],"domain_scores_gemma":[0.9895438,0.0062692408,0.0007811699,0.0016340672,0.0013610754,0.0004107263],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0050844178,0.0025473055,0.001597548,0.008401962,0.0011737931,0.0046291137,0.002431476,0.0023068925,0.010012686],"category_scores_gemma":[0.021320166,0.00095492415,0.0033888058,0.004641747,0.0010464196,0.0045186123,0.005504398,0.0030019174,0.009144573],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0019501406,0.0006071623,0.011769414,0.0033203454,0.001049447,0.0017168734,0.0012272592,0.020895982,0.04767389,0.028389163,0.1334812,0.7479191],"study_design_scores_gemma":[0.00033314346,0.0005063468,0.009795624,0.00085322664,0.0010580872,0.0017602434,0.0013846542,0.4933384,0.07416645,0.24099405,0.17550674,0.0003030511],"about_ca_topic_score_codex":0.004805256,"about_ca_topic_score_gemma":0.008380889,"teacher_disagreement_score":0.010012686,"about_ca_system_score_codex":0.0013714044,"about_ca_system_score_gemma":0.0037007614,"threshold_uncertainty_score":0.033495784},"labels":[],"label_agreement":null},{"id":"W7163361936","doi":"10.2196/86812","title":"A graph-based analysis approach for enhanced health study discoverability (Preprint)","year":2025,"lang":"en","type":"article","venue":"JMIR Medical Informatics","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"Discoverability; MEDLINE; Key (lock); Health informatics; Context (archaeology)","score_opus":0.01595976917289743,"score_gpt":0.3496653629922926,"score_spread":0.3337055938193952,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7163361936","genre_codex":"methods","genre_gemma":"empirical","domain_codex":null,"domain_gemma":"methods","model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":"methods","domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.023115799,0.0005512636,0.9561318,0.0016024252,0.00013979382,0.00048893027,0.0071606245,0.0070681954,0.0037411102],"genre_scores_gemma":[0.10526486,0.0003000901,0.88614655,0.00020135532,0.000082050625,0.00017669235,0.005145499,0.00048529592,0.002197585],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9982589,0.00054162054,0.00016507106,0.0004015512,0.00055382575,0.00007903121],"domain_scores_gemma":[0.9911522,0.0054766145,0.0006011587,0.0009328223,0.0015312202,0.00030597224],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.0024545991,0.00078072824,0.00064495835,0.014834812,0.0011985678,0.0032679571,0.001001961,0.0009379047,0.0075546377],"category_scores_gemma":[0.010998846,0.0004267791,0.0020941163,0.006287116,0.00063692196,0.002283206,0.0017578756,0.0012048027,0.0012598436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0005430808,0.00077973556,0.023743471,0.0017309312,0.0010462167,0.0012429003,0.0016323645,0.029293051,0.038849045,0.071513996,0.032299466,0.79732573],"study_design_scores_gemma":[0.00013994178,0.00021611569,0.016998988,0.00040485078,0.0009872719,0.0011131396,0.0011005824,0.68145233,0.018618807,0.22310607,0.05572484,0.00013710071],"about_ca_topic_score_codex":0.00927092,"about_ca_topic_score_gemma":0.017687468,"teacher_disagreement_score":0.9975454,"about_ca_system_score_codex":0.0009844885,"about_ca_system_score_gemma":0.0021110775,"threshold_uncertainty_score":0.025272846},"labels":[],"label_agreement":null},{"id":"W7164919706","doi":"10.1080/17153379.2010.12558959","title":"Otaku: Japan’s Database Animals.","year":2010,"lang":"en","type":"article","venue":"Pacific Affairs","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":true,"route_about_ca":false,"ca_institutions":"","funders":"","keywords":"","score_opus":0.011308245963998388,"score_gpt":0.2507373538871658,"score_spread":0.2394291079231674,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7164919706","genre_codex":"dataset","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.010794481,0.0030309057,0.034891773,0.0011974102,0.0003683042,0.0003070856,0.8335537,0.03577347,0.080082975],"genre_scores_gemma":[0.027252238,0.0026014987,0.03973975,0.00055611564,0.00011007117,0.00051474175,0.8993358,0.0037622768,0.026127515],"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9995969,0.000040192164,0.00008535037,0.00011130609,0.00011040684,0.00005589664],"domain_scores_gemma":[0.99901545,0.0001511456,0.00012007749,0.0002592761,0.00024888344,0.00020519097],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00066727377,0.0011737434,0.0008635482,0.0050048316,0.0009219702,0.0028808124,0.0014372015,0.000750033,0.06547169],"category_scores_gemma":[0.0021752373,0.00060589117,0.00042390678,0.008365864,0.00037586116,0.0039839214,0.0020976549,0.0010573021,0.038786724],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007739077,0.000118842865,0.006674267,0.0029362561,0.00012903434,0.0005306825,0.00069036905,0.0003179284,0.011113064,0.016473291,0.72414166,0.23610076],"study_design_scores_gemma":[0.00005748345,0.000019638932,0.005493816,0.00016387948,0.00009335042,0.00025470002,0.000309282,0.00091264653,0.0033778395,0.004096223,0.98518765,0.00003358299],"about_ca_topic_score_codex":0.010553615,"about_ca_topic_score_gemma":0.0116549,"teacher_disagreement_score":0.06547169,"about_ca_system_score_codex":0.0006435905,"about_ca_system_score_gemma":0.003151114,"threshold_uncertainty_score":0.2190246},"labels":[],"label_agreement":null},{"id":"W72549335","doi":"","title":"Automated Coding of Qualitative Interviews with Latent Semantic Analysis.","year":2007,"lang":"en","type":"article","venue":"WU Research","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":true,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"","funders":"York University","keywords":"Coding (social sciences); Computer science; Latent semantic analysis; Natural language processing; Qualitative research; Qualitative analysis; Psychology; Artificial intelligence; Data science; Sociology; Statistics; Mathematics","score_opus":0.2042462830884539,"score_gpt":0.5100896614593289,"score_spread":0.30584337837087494,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W72549335","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.062286273,0.00017063736,0.9115566,0.00079673895,0.00009363808,0.012125651,0.006648232,0.0014768834,0.004845339],"genre_scores_gemma":[0.13861383,0.00011117252,0.8352159,0.0000965915,0.000022614175,0.020368395,0.003622311,0.0001920116,0.001757138],"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.9481936,0.03988185,0.003835281,0.002764911,0.004668572,0.00065574853],"domain_scores_gemma":[0.83000433,0.11312803,0.013990448,0.0164719,0.025479728,0.00092565716],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04821673,0.0014572213,0.00079116505,0.008702178,0.002266852,0.0028310309,0.0019927912,0.0011060206,0.0070752357],"category_scores_gemma":[0.1158899,0.0006966581,0.0009896877,0.0083598215,0.0033864004,0.0034547877,0.0047768634,0.0013720453,0.0020815588],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0009407754,0.0003754663,0.0070939874,0.004373656,0.00009387907,0.0005306186,0.2457292,0.010927219,0.030730529,0.054149557,0.014707525,0.63034767],"study_design_scores_gemma":[0.00058664166,0.0005327476,0.029044513,0.0044254046,0.00015856109,0.0008814016,0.2678046,0.27290496,0.057452828,0.2298145,0.13577875,0.000615082],"about_ca_topic_score_codex":0.0039042747,"about_ca_topic_score_gemma":0.0061376495,"teacher_disagreement_score":0.04821673,"about_ca_system_score_codex":0.0054495824,"about_ca_system_score_gemma":0.007960929,"threshold_uncertainty_score":0.2549975},"labels":[],"label_agreement":null},{"id":"W73248818","doi":"10.1017/cbo9780511729942.019","title":"Comparison of replicate counts","year":2010,"lang":"en","type":"other","venue":"","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Achieve Life Sciences (Canada)","funders":"","keywords":"Replicate; Biology; Statistics; Mathematics","score_opus":0.025987823709722347,"score_gpt":0.35310610540689275,"score_spread":0.3271182816971704,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W73248818","genre_codex":"empirical","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6301232,0.0024307335,0.18053158,0.0007675226,0.000912033,0.00035468247,0.11620556,0.02258324,0.046091404],"genre_scores_gemma":[0.70591104,0.00059781095,0.12763183,0.00022915877,0.00015971214,0.0003957274,0.14002879,0.0035635715,0.021482337],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.99035746,0.0016774345,0.0009554089,0.0034058064,0.0031827532,0.00042102268],"domain_scores_gemma":[0.9588624,0.02091252,0.0018500796,0.011549806,0.0061231675,0.00070193113],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008683835,0.00088691333,0.0013221024,0.007795481,0.0010763207,0.0028536997,0.0017312437,0.0013547983,0.016635599],"category_scores_gemma":[0.04205279,0.0004451399,0.0017650982,0.007394742,0.00079836545,0.001966334,0.0014783533,0.0011611296,0.006225085],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0067404713,0.0010093177,0.1721691,0.003024007,0.004423214,0.001206563,0.0015047548,0.03131824,0.12867421,0.024742572,0.08111645,0.5440712],"study_design_scores_gemma":[0.0005138611,0.0012113621,0.3321997,0.00027642373,0.0024892932,0.0047268434,0.0023247884,0.16039334,0.284391,0.041686386,0.16924256,0.0005445162],"about_ca_topic_score_codex":0.0024932767,"about_ca_topic_score_gemma":0.004876565,"teacher_disagreement_score":0.016635599,"about_ca_system_score_codex":0.0010223332,"about_ca_system_score_gemma":0.0009651281,"threshold_uncertainty_score":0.055651605},"labels":[],"label_agreement":null},{"id":"W78175730","doi":"","title":"Examination of Routine Practice Patterns in the Hospital Information Data Warehouse: Use of OLAP and Rough Set Analysis with Clinician Feedback","year":2001,"lang":"en","type":"article","venue":"Europe PMC (PubMed Central)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Data warehouse; Online analytical processing; Test (biology); Interrogation; Computer science; Set (abstract data type); Database; Medicine; Medical record; Data mining; Medical emergency; Data science; Surgery","score_opus":0.03569251498221265,"score_gpt":0.26948332357815924,"score_spread":0.2337908085959466,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W78175730","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.6148173,0.0005504008,0.3621911,0.0023567544,0.00009358201,0.00288294,0.009308304,0.0034887865,0.0043108365],"genre_scores_gemma":[0.5963117,0.00019346492,0.39897648,0.000098661556,0.000027851336,0.0008758464,0.0030874799,0.00007048462,0.00035807813],"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","domain_scores_codex":[0.9851147,0.00798326,0.0026133338,0.0010036569,0.003069076,0.00021607439],"domain_scores_gemma":[0.8998282,0.07586198,0.006599985,0.006526783,0.010395153,0.0007879751],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02250321,0.000655542,0.0012848636,0.0077961227,0.0006867785,0.003722764,0.0009323938,0.00047027427,0.0010032334],"category_scores_gemma":[0.08814341,0.00044201175,0.0009643408,0.008872047,0.0005283572,0.0027230668,0.0011098159,0.0008939472,0.00032892966],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.002242983,0.0012014547,0.21603791,0.0022823836,0.00088883535,0.00081776583,0.012864509,0.04225576,0.007476215,0.004665446,0.00605326,0.7032136],"study_design_scores_gemma":[0.000482168,0.002292446,0.19466694,0.0007449705,0.0006947193,0.0010373804,0.01504527,0.73075074,0.017686004,0.025344446,0.010817514,0.00043728406],"about_ca_topic_score_codex":0.0057874196,"about_ca_topic_score_gemma":0.005205435,"teacher_disagreement_score":0.02250321,"about_ca_system_score_codex":0.0014385874,"about_ca_system_score_gemma":0.0023937102,"threshold_uncertainty_score":0.11900979},"labels":[],"label_agreement":null},{"id":"W801626237","doi":"","title":"A phylogenomic falsification of the chromalveolate hypothesis","year":2010,"lang":"en","type":"article","venue":"Open Repository and Bibliography (University of Liège)","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Dalhousie University; Université de Montréal","funders":"","keywords":"Computer science","score_opus":0.012189659689056451,"score_gpt":0.20810403095094132,"score_spread":0.19591437126188488,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W801626237","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.70464164,0.005021203,0.18672715,0.031185646,0.0022947625,0.00012506059,0.008966401,0.0019807727,0.059057362],"genre_scores_gemma":[0.9572696,0.001207587,0.031800635,0.0024773113,0.0003115572,0.000030539904,0.0036483693,0.00017862397,0.0030757014],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","domain_scores_codex":[0.99855536,0.00033916454,0.000102551414,0.0004693693,0.00036688175,0.0001665771],"domain_scores_gemma":[0.9861557,0.008476056,0.0012354249,0.0027647302,0.000845721,0.00052241905],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0033031115,0.0005798388,0.00090154505,0.0033798988,0.0020743052,0.0020904378,0.0020894413,0.001852861,0.009443464],"category_scores_gemma":[0.012796481,0.00021938284,0.000887156,0.003373384,0.0029228244,0.0021602816,0.001608395,0.0019364092,0.0011722913],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0035962153,0.00043264727,0.087963335,0.0014980815,0.00091649883,0.027418103,0.0032093304,0.008028325,0.16721477,0.45360386,0.01821329,0.2279055],"study_design_scores_gemma":[0.0005627725,0.00042888912,0.06649687,0.00082253816,0.0011516977,0.021653086,0.0040611657,0.035711206,0.09586374,0.6095358,0.16345678,0.0002555064],"about_ca_topic_score_codex":0.0010395374,"about_ca_topic_score_gemma":0.0010545147,"teacher_disagreement_score":0.009443464,"about_ca_system_score_codex":0.0006378505,"about_ca_system_score_gemma":0.0010029189,"threshold_uncertainty_score":0.031591535},"labels":[{"model":"gemma","categories":[],"domain":null,"study_design":"theoretical_or_conceptual","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"},{"model":"gpt","categories":[],"domain":null,"study_design":"design_other","genre":"empirical","about_ca_system":false,"about_ca_topic":false,"confidence":"low"}],"label_agreement":"split"},{"id":"W91636976","doi":"10.1007/978-3-642-30284-8_50","title":"Evaluating Scientific Hypotheses Using the SPARQL Inferencing Notation","year":2012,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":12,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Carleton University","funders":"","keywords":"SPARQL; Notation; Computer science; Domain (mathematical analysis); TRACE (psycholinguistics); Flexibility (engineering); Task (project management); Data science; RDF; Information retrieval; Semantic Web; Mathematics","score_opus":0.10815389168132306,"score_gpt":0.3470901800926138,"score_spread":0.23893628841129072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W91636976","genre_codex":"methods","genre_gemma":"methods","domain_codex":null,"domain_gemma":null,"model_version":"metacan-v3-hybrid-931329e0061c","genre_candidate":"methods","genre_consensus":"methods","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.013336168,0.0004741633,0.9607543,0.0010928665,0.00016327585,0.00035420092,0.003915699,0.013462547,0.0064468062],"genre_scores_gemma":[0.12488615,0.00073316705,0.8574164,0.00058776577,0.00014987712,0.00031674188,0.011038728,0.0012995276,0.00357168],"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","domain_scores_codex":[0.9926099,0.0024801083,0.0009897198,0.0008497122,0.0028157057,0.0002548281],"domain_scores_gemma":[0.9842594,0.011629051,0.00065912027,0.0017359147,0.0014663651,0.000250134],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009920328,0.0013268327,0.001342535,0.003818566,0.0010665221,0.006571533,0.00348314,0.0016669505,0.01434638],"category_scores_gemma":[0.026525669,0.0009931296,0.00330589,0.0027685778,0.0013579781,0.007778368,0.0044280565,0.0021020523,0.0032387467],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0007647148,0.00064533,0.0059229117,0.0021212946,0.00058250467,0.0018641695,0.00101481,0.07234555,0.01911372,0.24563143,0.03739904,0.61259454],"study_design_scores_gemma":[0.00020420746,0.00019817,0.0010159336,0.00039705486,0.00042808434,0.0007784647,0.0004960228,0.42092046,0.04628733,0.4785673,0.050581966,0.00012504972],"about_ca_topic_score_codex":0.002591936,"about_ca_topic_score_gemma":0.004332396,"teacher_disagreement_score":0.01434638,"about_ca_system_score_codex":0.0009603309,"about_ca_system_score_gemma":0.002007413,"threshold_uncertainty_score":0.052464366},"labels":[],"label_agreement":null}]}