{"id":"W2935566373","doi":"10.1002/jrs.5608","title":"Optimized preprocessing and machine learning for quantitative Raman spectroscopy in biology","year":2019,"lang":"en","type":"preprint","venue":"Journal of Raman Spectroscopy","topic":"Spectroscopy Techniques in Biomedical and Chemical Research","field":"Biochemistry, Genetics and Molecular Biology","cited_by":2,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"Semiconductor Research Corporation","keywords":"Raman spectroscopy; Robustness (evolution); Preprocessor; Computer science; Artificial intelligence; Medical diagnosis; Machine learning; Process (computing); Set (abstract data type); Biological system; Optics; Biology; Physics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.001444462,0.0004921253,0.00107933,0.0003924984,0.00009030963,0.0001324435,0.0007471741,0.0008250169,0.0000506113],"category_scores_gemma":[0.001175142,0.0004209362,0.0003530635,0.0001740131,0.0004099571,0.00001337896,0.0007578746,0.00198267,0.000001856517],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0002035513,"about_ca_system_score_gemma":0.0004634335,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00002758191,"about_ca_topic_score_gemma":0.000008046295,"domain_scores_codex":[0.9968647,0.0002192595,0.001006858,0.0008328446,0.0003615758,0.0007147264],"domain_scores_gemma":[0.9978998,0.0002105733,0.0009220297,0.0004284024,0.000302005,0.0002371361],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.003180804,0.0002113114,0.004549421,0.0004427508,0.0002029914,0.00001434421,0.0001097145,0.0001951318,0.9890049,0.0004039273,0.001442962,0.0002418086],"study_design_scores_gemma":[0.003256862,0.004469892,0.0002256894,0.0005872822,0.00008384211,0.00006723829,0.0001167059,0.003849049,0.9637035,0.01393284,0.009138445,0.000568629],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7056502,0.01387592,0.2766354,0.000835856,0.0005595234,0.001060979,0.00005624428,0.00003302453,0.001292859],"genre_scores_gemma":[0.6411905,0.01238357,0.3440245,0.0002451983,0.0009629571,0.00009275824,0.0002998354,0.000141083,0.0006595483],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.06738914,"threshold_uncertainty_score":0.9998242,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02545766168058384,"score_gpt":0.3902677291017067,"score_spread":0.3648100674211229,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}