{"id":"W4309373803","doi":"10.1101/2022.11.18.517107","title":"Predicting environmental stressor levels with machine learning: a comparison between amplicon sequencing, metagenomics, and total RNA sequencing based on taxonomically assigned data","year":2022,"lang":"en","type":"preprint","venue":"bioRxiv (Cold Spring Harbor Laboratory)","topic":"Microbial Community Ecology and Physiology","field":"Environmental Science","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Guelph","funders":"","keywords":"Metagenomics; Amplicon sequencing; Amplicon; Biology; Context (archaeology); Deep sequencing; Massive parallel sequencing; Machine learning; Computational biology; Artificial intelligence; DNA sequencing; Computer science; Genetics; Genome; 16S ribosomal RNA; Polymerase chain reaction; Gene","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006997444,0.00137068,0.0009055621,0.002163468,0.0003398107,0.001236982,0.0007686228,0.0009429923,0.0003506107],"category_scores_gemma":[0.007917383,0.0003299878,0.00117925,0.001279833,0.0004565417,0.001185787,0.0007452349,0.0007054124,0.0001793454],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005990466,"about_ca_system_score_gemma":0.0005086908,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001466456,"about_ca_topic_score_gemma":0.001663573,"domain_scores_codex":[0.9971757,0.001429525,0.0002336468,0.0006021482,0.0004474687,0.0001114973],"domain_scores_gemma":[0.993663,0.004848978,0.000491484,0.0002950998,0.0005589326,0.0001424638],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.003095338,0.001370334,0.2707416,0.0009400719,0.00250549,0.0001854443,0.0003631848,0.2917688,0.05403734,0.0008511203,0.0006493298,0.3734918],"study_design_scores_gemma":[0.00005708132,0.001544395,0.06177584,0.0001071296,0.0002904299,0.00009665771,0.000170367,0.9126513,0.0208401,0.001749979,0.0006430197,0.0000736914],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9223721,0.003178015,0.07211143,0.0003441635,0.00005657586,0.0001195322,0.0004582433,0.0003553784,0.001004498],"genre_scores_gemma":[0.932695,0.0008170545,0.06536973,0.0001105908,0.0000410087,0.00009067039,0.0006606582,0.00004126252,0.0001738674],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.006997444,"threshold_uncertainty_score":0.0370065,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04492372406273521,"score_gpt":0.2333324514609866,"score_spread":0.1884087273982514,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}