{"id":"W4412396400","doi":"10.2196/80519","title":"Correction: Large Language Model–Assisted Risk-of-Bias Assessment in Randomized Controlled Trials Using the Revised Risk-of-Bias Tool: Evaluation Study","year":2025,"lang":"en","type":"erratum","venue":"Journal of Medical Internet Research","topic":"Radiomics and Machine Learning in Medical Imaging","field":"Medicine","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"JMIR Publications","funders":"","keywords":"Randomized controlled trial; Computer science; Risk assessment; Psychology; Data mining; Medicine; Medical physics; Computer security","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","research_integrity","insufficient_payload"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.4366267,0.0004605018,0.009284034,0.002033093,0.0001351862,0.0001436444,0.001172622,0.0008157206,0.001318775],"category_scores_gemma":[0.543495,0.0002336809,0.002459004,0.0009860876,0.0005774109,0.00009269111,0.0004108595,0.01585904,0.000001277705],"about_ca_system_candidate":true,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008495175,"about_ca_system_score_gemma":0.009246351,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002182932,"about_ca_topic_score_gemma":0.0001762389,"domain_scores_codex":[0.9135394,0.06425432,0.007808615,0.0005384536,0.01318315,0.0006760373],"domain_scores_gemma":[0.9315124,0.0536124,0.009156789,0.0009669448,0.004265503,0.0004859731],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.1785929,0.00553058,0.008669086,0.001509205,0.01781071,0.0009171056,0.006819474,0.003012744,0.00007654206,0.0001126433,0.7073702,0.06957882],"study_design_scores_gemma":[0.2236755,0.0006331802,0.0008084091,0.009938342,0.005301963,0.00007973383,0.002462424,0.7554946,0.000009065702,0.0001174714,0.001350686,0.0001286689],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6774113,0.05860822,0.1257651,0.01736652,0.06924327,0.03109838,0.000107956,0.0000661912,0.02033305],"genre_scores_gemma":[0.9401432,0.01710379,0.001346121,0.0002482098,0.004440302,0.0002682477,0.0001020396,0.0001048596,0.03624318],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.7524819,"threshold_uncertainty_score":0.9995942,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1551637847947682,"score_gpt":0.5316779197303239,"score_spread":0.3765141349355556,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}