{"meta":{"page":1,"per_page":50,"max_per_page":100,"total":7,"total_is_capped":false,"direct_labels_cover":0,"predictions_cover":7,"direct_label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline (scores rank; they never assert a category)","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12","author_layer_release":"2026-06-26"},"query_hash":"805c37b54ad1","filters":{"venue":"Biostatistics & Epidemiology"}},"results":[{"id":"W2971936035","doi":"10.1080/24709360.2019.1660110","title":"A frequentist mixture modeling of stop signal reaction times","year":2019,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Blind Source Separation Techniques","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Hospital for Sick Children; Public Health Ontario; University of Toronto","funders":"","keywords":"Frequentist inference; SIGNAL (programming language); Computer science; Econometrics; Statistics; Algorithm; Mathematics; Bayesian probability; Bayesian inference","authors":[{"name":"Mohsen Soltanifar","is_ca":true},{"name":"Annie Dupuis","is_ca":true},{"name":"Russell Schachar","is_ca":true},{"name":"Michael Escobar","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.03552478450841576,"gpt":0.3131540366445244,"spread":0.2776292521361086,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005596126,0.001293364,0.002230532,0.001884073,0.0007614632,0.002145903,0.004007662,0.002362171,0.004400888],"category_scores_gemma":[0.01266072,0.001192961,0.003189615,0.001722767,0.00117793,0.002039211,0.001587838,0.002716473,0.001310977],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001134198,"about_ca_system_score_gemma":0.0009955291,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01139045,"about_ca_topic_score_gemma":0.00931157,"domain_scores_codex":[0.9973907,0.001052489,0.0001271509,0.0008041238,0.0003557783,0.0002698677],"domain_scores_gemma":[0.9953342,0.003155971,0.000431488,0.0004725727,0.0004724222,0.0001332524],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001003771,0.0003456686,0.01443559,0.0002479641,0.0007903553,0.0005104321,0.0009724081,0.7607659,0.004917713,0.1291677,0.004049176,0.08279334],"study_design_scores_gemma":[0.00002536741,0.00004594873,0.001633949,0.00001060108,0.00005083791,0.00005868609,0.00001998195,0.9821177,0.0001788396,0.01520825,0.0006193136,0.00003050707],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.10342,0.0004629641,0.8913403,0.0006479516,0.0001251345,0.0002493758,0.000841113,0.0006911856,0.002221846],"genre_scores_gemma":[0.7969844,0.0006306889,0.1810214,0.0003021797,0.0002374776,0.000983448,0.002373471,0.0003019219,0.01716499],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01139045,"threshold_uncertainty_score":0.02959549,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W3047938449","doi":"10.1080/24709360.2020.1796176","title":"Number needed to test: quantifying risk stratification provided by diagnostic tests and risk predictions","year":2020,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"BRCA gene mutations in cancer","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":true,"ca_venue":false,"about_ca":false},"ca_institutions":"McGill University","funders":"Division of Cancer Epidemiology and Genetics, National Cancer Institute; Fonds de Recherche du Québec-Société et Culture; Canadian Network for Research and Innovation in Machining Technology, Natural Sciences and Engineering Research Council of Canada","keywords":"Risk stratification; Genetic testing; Medicine; Risk assessment; Psychological intervention; Disease; Metric (unit); Test (biology); Actuarial science; Econometrics; Computer science; Risk analysis (engineering); Mathematics; Internal medicine; Economics; Biology; Operations management","authors":[{"name":"Hormuzd A. Katki","is_ca":false},{"name":"Rajib Dey","is_ca":true},{"name":"Paramita Saha‐Chaudhuri","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.04071876991350831,"gpt":0.3399215882456961,"spread":0.2992028183321878,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.02008147,0.001647288,0.001258041,0.00301107,0.0007525259,0.003530249,0.001382419,0.003064234,0.006142885],"category_scores_gemma":[0.1620321,0.0006732614,0.001584668,0.001646882,0.002472046,0.005527728,0.002465329,0.0027902,0.0008566203],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002375586,"about_ca_system_score_gemma":0.001328013,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004804189,"about_ca_topic_score_gemma":0.002751273,"domain_scores_codex":[0.9858382,0.007525933,0.0007243167,0.001866149,0.0033981,0.0006472835],"domain_scores_gemma":[0.835147,0.1421274,0.00864901,0.007415411,0.005114721,0.00154641],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.003164317,0.0004311723,0.2315292,0.0007160156,0.0009742102,0.0005986557,0.0009209916,0.504871,0.003185672,0.05062169,0.006138662,0.1968483],"study_design_scores_gemma":[0.0002526688,0.001297075,0.07388232,0.0005414124,0.0004992886,0.0015803,0.0006777217,0.6962633,0.00600891,0.209235,0.009398481,0.0003635795],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"methods","genre_scores_codex":[0.5512115,0.006105016,0.38415,0.01379155,0.0005687849,0.0004296355,0.004604788,0.002214751,0.03692401],"genre_scores_gemma":[0.9478722,0.000444475,0.04845377,0.0006685436,0.0001148374,0.0001464602,0.0009821423,0.0001633689,0.001154215],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.02008147,"threshold_uncertainty_score":0.1062022,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W4280498397","doi":"10.1080/24709360.2022.2069457","title":"Doubly weighted estimating equations and weighted multiple imputation for causal inference with an incomplete subgroup variable","year":2022,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Advanced Causal Inference Techniques","field":"Mathematics","cited_by":1,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":true,"ca_venue":false,"about_ca":false},"ca_institutions":"London Health Sciences Centre; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada; Canadian Institutes of Health Research","keywords":"Missing data; Statistics; Observational study; Mathematics; Estimator; Causal inference; Inference; Imputation (statistics); Subgroup analysis; Variable (mathematics); Propensity score matching; Econometrics; Medicine; Computer science; Artificial intelligence","authors":[{"name":"Meaghan S. Cuerden","is_ca":true},{"name":"Liqun Diao","is_ca":true},{"name":"Cecilia A. Cotton","is_ca":true},{"name":"Richard J. Cook","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1809378732533116,"gpt":0.4283506853874128,"spread":0.2474128121341011,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03717695,0.00149891,0.002848067,0.003439223,0.0006042957,0.001797065,0.004506961,0.002216105,0.005063281],"category_scores_gemma":[0.1340474,0.00118277,0.003379401,0.005955017,0.001353273,0.00296492,0.002973199,0.004042274,0.001121609],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001272902,"about_ca_system_score_gemma":0.002965702,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005871609,"about_ca_topic_score_gemma":0.005448453,"domain_scores_codex":[0.9631482,0.03198755,0.001091324,0.001553288,0.00189044,0.0003291885],"domain_scores_gemma":[0.9233024,0.06291646,0.004631321,0.006668874,0.002225214,0.0002556799],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0001804425,0.0001351808,0.006998374,0.0007146808,0.001790164,0.0003866144,0.0005115935,0.1800976,0.0005517682,0.6010912,0.004110471,0.2034319],"study_design_scores_gemma":[0.0001288506,0.0001060724,0.001095151,0.0001759827,0.0002474633,0.0001463604,0.00005149532,0.4723447,0.0003892964,0.5189145,0.006339387,0.00006070873],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.0009204439,0.0002375117,0.9982244,0.0001565048,0.0000244609,0.0000664773,0.0001103399,0.00007766005,0.0001822518],"genre_scores_gemma":[0.05255783,0.001207951,0.9428252,0.0002064416,0.0001504809,0.001144268,0.0006888455,0.00008945382,0.00112948],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.03717695,"threshold_uncertainty_score":0.1966128,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W4406596694","doi":"10.1080/24709360.2025.2451519","title":"Finite Markov chains with absorbing states and mis-specified random effects: application to cognitive data","year":2025,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Bayesian Methods and Mixture Models","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Waterloo","funders":"National Center for Advancing Translational Sciences; National Institute on Aging","keywords":"Markov chain; Statistical physics; Cognition; Computer science; Mathematics; Physics; Psychology; Statistics","authors":[{"name":"Pei Wang","is_ca":false},{"name":"Changrui Liu","is_ca":false},{"name":"Jiyeon Park","is_ca":false},{"name":"Suzanne L. Tyas","is_ca":true},{"name":"Richard J. Kryscio","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.03042029286051502,"gpt":0.3396504909591254,"spread":0.3092301980986104,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.03303747,0.0009430817,0.002018043,0.002994151,0.001346087,0.002548339,0.002562837,0.002399084,0.003393284],"category_scores_gemma":[0.1231917,0.001063542,0.002569143,0.002922888,0.003014392,0.003007962,0.002717815,0.003780801,0.000407288],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002139388,"about_ca_system_score_gemma":0.00275303,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01562281,"about_ca_topic_score_gemma":0.01463199,"domain_scores_codex":[0.9888973,0.008650241,0.0003827405,0.001001984,0.0007494675,0.0003181649],"domain_scores_gemma":[0.7983513,0.1874229,0.005523313,0.004957657,0.002986913,0.0007577955],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0002028284,0.0001711773,0.01704195,0.0003084547,0.0003727096,0.0007619816,0.001432063,0.4874779,0.0006093101,0.4224035,0.001498698,0.06771949],"study_design_scores_gemma":[0.00003078807,0.00003445231,0.001268609,0.00007567551,0.000048158,0.0001043887,0.0000998151,0.7728357,0.0001936713,0.2242497,0.001022455,0.00003657935],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01840899,0.0002243716,0.9799269,0.0003460363,0.00002553935,0.0001190693,0.0001220039,0.0002114572,0.0006158228],"genre_scores_gemma":[0.4339052,0.0007808593,0.5614521,0.0003602799,0.0001335846,0.0009444561,0.0005368321,0.0001385735,0.001748],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.03303747,"threshold_uncertainty_score":0.1747209,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2760713401","doi":"10.1080/24709360.2017.1359356","title":"Analysis of progressive multi-state models with misclassified states: likelihood and pairwise likelihood methods","year":2017,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Statistical Methods and Bayesian Inference","field":"Mathematics","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":true,"ca_venue":false,"about_ca":false},"ca_institutions":"Western University; University of Waterloo","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Pairwise comparison; Inference; Computer science; Maximum likelihood; State (computer science); Data mining; Artificial intelligence; Machine learning; Econometrics; Statistics; Mathematics; Algorithm","authors":[{"name":"Grace Y. Yi","is_ca":true},{"name":"Wenqing He","is_ca":true},{"name":"Feng He","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1924233008744135,"gpt":0.4705267785926885,"spread":0.2781034777182749,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05280696,0.001758306,0.002934635,0.002459469,0.0008251612,0.002660493,0.005806346,0.002661049,0.002932034],"category_scores_gemma":[0.1655959,0.001311409,0.003658987,0.002459363,0.003498622,0.004598537,0.004185762,0.005258273,0.0004156089],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001503042,"about_ca_system_score_gemma":0.002454973,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004188429,"about_ca_topic_score_gemma":0.003122053,"domain_scores_codex":[0.980297,0.01484207,0.0006341262,0.002339755,0.001519701,0.0003675309],"domain_scores_gemma":[0.8183271,0.164987,0.00723478,0.006191319,0.002301898,0.0009578791],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002430041,0.0001530258,0.01690692,0.0003879428,0.001173562,0.0006275011,0.0008359479,0.5736148,0.0005792502,0.3069158,0.001375738,0.09718654],"study_design_scores_gemma":[0.00004109214,0.00008237202,0.001529767,0.00007576086,0.0001112735,0.0001627779,0.00006169231,0.7792315,0.0003060292,0.2173722,0.0009768215,0.0000487371],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.006503072,0.0002894866,0.9925913,0.0002691098,0.00002318289,0.00004949633,0.00005669313,0.00004890522,0.0001689105],"genre_scores_gemma":[0.4363988,0.001475685,0.5564958,0.0003396647,0.0003084788,0.0008701027,0.000826754,0.0001609761,0.003123759],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.05280696,"threshold_uncertainty_score":0.2792732,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2997098473","doi":"10.1080/24709360.2019.1699341","title":"Cohort study design for illness-death processes with disease status under intermittent observation","year":2019,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Statistical Methods and Inference","field":"Mathematics","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Waterloo; University of Toronto","funders":"","keywords":"Censoring (clinical trials); Cohort; Medicine; Cohort study; Sample size determination; Prospective cohort study; Disease; Incidence (geometry); Confidence interval; Event (particle physics); Demography; Statistics; Internal medicine; Mathematics; Pathology","authors":[{"name":"Nathalie C. Moon","is_ca":true},{"name":"Leilei Zeng","is_ca":true},{"name":"Richard J. Cook","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.2551980163011247,"gpt":0.4321040283620619,"spread":0.1769060120609373,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.07755016,0.001317899,0.002076796,0.001835834,0.001106388,0.001840505,0.004098156,0.002901903,0.009417197],"category_scores_gemma":[0.09129766,0.0007347465,0.001977705,0.002921654,0.001868829,0.002194629,0.003060699,0.003408845,0.001262329],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001184315,"about_ca_system_score_gemma":0.003648924,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00173269,"about_ca_topic_score_gemma":0.001876528,"domain_scores_codex":[0.9525067,0.03900289,0.001545265,0.004010374,0.002298478,0.0006363947],"domain_scores_gemma":[0.9403328,0.04382235,0.004981468,0.007854248,0.002245286,0.0007639326],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.002692686,0.0007687104,0.03682445,0.002716593,0.002358316,0.0007585919,0.001812574,0.02996368,0.002295281,0.8005115,0.01115961,0.1081381],"study_design_scores_gemma":[0.004083075,0.01101442,0.01609308,0.001475821,0.002794628,0.001489763,0.0009673047,0.3315643,0.00179712,0.5340964,0.09431005,0.0003140598],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01571208,0.001645756,0.9673935,0.001054461,0.001142911,0.008871267,0.001899638,0.0002498286,0.002030491],"genre_scores_gemma":[0.1922992,0.004570465,0.7145315,0.001615741,0.001187915,0.07598013,0.002464838,0.0001027088,0.007247488],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.9224498,"threshold_uncertainty_score":0.4101293,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W4413019256","doi":"10.1080/24709360.2025.2539589","title":"Selection of mediators and dependence structure for high-dimensional mediation analysis","year":2025,"lang":"en","type":"article","venue":"Biostatistics & Epidemiology","topic":"Child and Adolescent Psychosocial and Emotional Development","field":"Psychology","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Waterloo","funders":"","keywords":"Mediation; Psychology; Selection (genetic algorithm); Computer science; Artificial intelligence; Political science","authors":[{"name":"Lijia Wang","is_ca":true},{"name":"Yeying Zhu","is_ca":true},{"name":"Richard J. Cook","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.01788345427079987,"gpt":0.3335996795926779,"spread":0.315716225321878,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.04233093,0.001339738,0.002399951,0.001497993,0.001750173,0.002320038,0.002909447,0.001808458,0.00811728],"category_scores_gemma":[0.08709102,0.0009670762,0.002804015,0.001657492,0.002158391,0.002712152,0.005412834,0.004433571,0.0008844887],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007035938,"about_ca_system_score_gemma":0.004118843,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001518602,"about_ca_topic_score_gemma":0.001921103,"domain_scores_codex":[0.9678361,0.02854282,0.0006332655,0.001349815,0.001123484,0.0005144237],"domain_scores_gemma":[0.947125,0.04392151,0.002173558,0.004422078,0.001747443,0.0006104553],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.001358049,0.001537677,0.05160227,0.001105731,0.002000067,0.001342639,0.003156923,0.1583259,0.006904089,0.5394608,0.007000891,0.2262049],"study_design_scores_gemma":[0.0002361373,0.0004275005,0.005954145,0.0001487619,0.0002338755,0.0001719029,0.0003327226,0.7750584,0.001719025,0.2120308,0.003594622,0.00009215392],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.0113613,0.00008926803,0.9870703,0.0003689849,0.00002123286,0.0003653445,0.00009664583,0.0001399427,0.0004870245],"genre_scores_gemma":[0.2611106,0.000282519,0.7329424,0.000267852,0.00009179681,0.003015399,0.0005933368,0.0001166248,0.001579535],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.04233093,"threshold_uncertainty_score":0.22387,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null}]}