{"id":"W4223653240","doi":"10.1111/1755-0998.13615","title":"A setback into a success: What can batch effects tell us about best practices in genomics?","year":2022,"lang":"en","type":"article","venue":"Molecular Ecology Resources","topic":"Genetic diversity and population structure","field":"Biochemistry, Genetics and Molecular Biology","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"Université Laval","funders":"","keywords":"Population genomics; Biology; Genomics; DNA sequencing; Data science; Inference; Population; Causal inference; Whole genome sequencing; Ecology; Computational biology; Genome; Genetics; Computer science; Artificial intelligence","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.0002691364,0.0002164175,0.0002422807,0.0001391012,0.0003265584,0.00009840676,0.0005004199,0.0002184752,0.0001572965],"category_scores_gemma":[0.0001297137,0.0002494868,0.0001005872,0.0002145239,0.0001096571,0.00001123436,0.0006435533,0.000291917,0.00001757375],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0000678369,"about_ca_system_score_gemma":0.0001029378,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0005216628,"about_ca_topic_score_gemma":0.002907666,"domain_scores_codex":[0.9981364,0.0004892661,0.000238195,0.0005802243,0.000188832,0.0003670664],"domain_scores_gemma":[0.9991122,0.00004452227,0.0002999032,0.0003975103,0.00004526747,0.0001006087],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0004449449,0.0003366381,0.6090846,0.0001447132,0.0002948282,0.0005692212,0.004743741,0.02270994,0.3570575,0.00005407816,0.0009532637,0.003606582],"study_design_scores_gemma":[0.003649513,0.002021691,0.2535142,0.00003460144,0.0002079465,0.000298577,0.005746211,0.0003093155,0.08093036,0.0004331356,0.6516855,0.001168852],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.994611,0.003482065,0.00002129104,0.0006692502,0.0003436836,0.0003710344,0.00001666383,0.00001183746,0.0004731488],"genre_scores_gemma":[0.9964216,0.0002441407,0.0003984985,0.002076468,0.00007155963,0.0000728678,0.0001867469,0.00002654054,0.0005015326],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.6507323,"threshold_uncertainty_score":0.9999957,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.007858059042478608,"score_gpt":0.2449686423113842,"score_spread":0.2371105832689056,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}