{"id":"W4210520314","doi":"10.1016/j.ijrobp.2022.01.050","title":"Comprehensive Quantitative Evaluation of Variability in Magnetic Resonance-Guided Delineation of Oropharyngeal Gross Tumor Volumes and High-Risk Clinical Target Volumes: An R-IDEAL Stage 0 Prospective Study","year":2022,"lang":"en","type":"article","venue":"International Journal of Radiation Oncology*Biology*Physics","topic":"Head and Neck Cancer Studies","field":"Medicine","cited_by":30,"is_retracted":false,"has_abstract":false,"ca_institutions":"Health Sciences Centre; Sunnybrook Health Science Centre","funders":"National Institute of Dental and Craniofacial Research; National Institute of Biomedical Imaging and Bioengineering; National Cancer Institute","keywords":"Medicine; Magnetic resonance imaging; Head and neck cancer; Nuclear medicine; Stage (stratigraphy); Radiology; Institutional review board; Radiation therapy; Gold standard (test); Medical physics; Surgery","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00568771,0.0001788777,0.0009706989,0.0002076298,0.00009878109,0.00000725175,0.0002069013,0.00008801914,0.00008363932],"category_scores_gemma":[0.002097624,0.0001676268,0.0001282119,0.0002773447,0.0003511673,0.0002365405,0.0001022792,0.0005922955,6.859079e-7],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008635464,"about_ca_system_score_gemma":0.001109002,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0004723974,"about_ca_topic_score_gemma":0.00008175633,"domain_scores_codex":[0.9938505,0.00326263,0.001661868,0.0003558049,0.0007032955,0.0001659524],"domain_scores_gemma":[0.9936872,0.001219514,0.001757136,0.0001610193,0.003095122,0.00008007351],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.002423282,0.002277202,0.9063793,0.00001778133,0.0003156209,0.00001800233,0.003689914,0.002932909,0.0005113583,0.001175403,0.00008592344,0.08017328],"study_design_scores_gemma":[0.01366277,0.01399981,0.8896207,0.00003137205,0.0002634824,0.00002320243,0.003320933,0.07406871,0.0001770127,0.004272745,0.0004400567,0.0001192247],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9943961,0.001441666,0.001529991,0.0002893355,0.001005046,0.0009716275,0.0003164135,0.00001038675,0.00003944038],"genre_scores_gemma":[0.9951968,0.0002805206,0.003865219,0.0001224952,0.0003751384,0.00006420886,0.00007007802,0.00001664629,0.000008963334],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.08005406,"threshold_uncertainty_score":0.6835624,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.06657668158867787,"score_gpt":0.4227288227355428,"score_spread":0.3561521411468649,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}