{"id":"W2587370908","doi":"10.1093/annonc/mdw387.02","title":"Do contemporary randomized controlled trials meet ESMO thresholds for clinically meaningful benefit?","year":2016,"lang":"en","type":"article","venue":"Annals of Oncology","topic":"Health Systems, Economic Evaluations, Quality of Life","field":"Economics, Econometrics and Finance","cited_by":4,"is_retracted":false,"has_abstract":false,"ca_institutions":"Princess Margaret Cancer Centre","funders":"","keywords":"Medicine; Randomized controlled trial; Internal medicine; Clinical endpoint; Breast cancer; Sample size determination; Oncology; Surrogate endpoint; Clinical trial; Cancer","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.3488616,0.001612457,0.01266838,0.003919868,0.00120732,0.01069714,0.005011221,0.009304472,0.01118042],"category_scores_gemma":[0.6739094,0.001834793,0.008601801,0.005003179,0.00702614,0.01721996,0.004742037,0.01182839,0.001965245],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.005661492,"about_ca_system_score_gemma":0.006308997,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001737738,"about_ca_topic_score_gemma":0.003006706,"domain_scores_codex":[0.6728052,0.2299204,0.05847561,0.0120172,0.02336177,0.00341988],"domain_scores_gemma":[0.1888082,0.7359479,0.03981933,0.01390573,0.0181861,0.00333263],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.05570738,0.0006892825,0.02409183,0.1395772,0.05867568,0.0003137197,0.002067772,0.004694112,0.001388864,0.125438,0.06282385,0.5245324],"study_design_scores_gemma":[0.07179362,0.009929507,0.04325091,0.1097015,0.05836653,0.001006438,0.002339332,0.01409963,0.002077911,0.4429073,0.2439608,0.0005664263],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"review","genre_gemma":"empirical","genre_scores_codex":[0.02660123,0.6322075,0.04879585,0.2463917,0.02169608,0.00361916,0.002585311,0.0004323387,0.01767082],"genre_scores_gemma":[0.5553796,0.0866969,0.07193593,0.2533261,0.02148882,0.006930714,0.002484368,0.0003996979,0.001357791],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.6511384,"threshold_uncertainty_score":0.8029695,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.7286972974741612,"score_gpt":0.5681810570550706,"score_spread":0.1605162404190906,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}