{"id":"W2970783423","doi":"10.1177/0272989x19856617","title":"Exclusion Criteria as Measurements I: Identifying Invalid Responses","year":2019,"lang":"en","type":"article","venue":"Medical Decision Making","topic":"Health Systems, Economic Evaluations, Quality of Life","field":"Economics, Econometrics and Finance","cited_by":28,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Social Sciences and Humanities Research Council of Canada; National Institutes of Health; National Center for Advancing Translational Sciences; Riksbankens Jubileumsfond","keywords":"Valuation (finance); Sample (material); Psychology; Actuarial science; Econometrics; Medicine; Social psychology; Economics","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","metaepi_narrow","insufficient_payload"],"consensus_categories":["metaresearch","insufficient_payload"],"category_scores_codex":[0.03932392,0.0002277709,0.0009395726,0.0005507794,0.0002777802,0.000218415,0.0006737682,0.0002966196,0.02365913],"category_scores_gemma":[0.03994687,0.0002605151,0.0001691978,0.0003011823,0.00005738528,0.0005829884,0.0003577874,0.0003132476,0.02828978],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005035382,"about_ca_system_score_gemma":0.0002353729,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0001310729,"about_ca_topic_score_gemma":0.00003575163,"domain_scores_codex":[0.9930906,0.0006307467,0.004106145,0.0008213972,0.0008287204,0.0005223579],"domain_scores_gemma":[0.9941412,0.003269306,0.001349293,0.0007808554,0.0001040402,0.0003552997],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"not_applicable","study_design_scores_codex":[0.001249379,0.0007658907,0.5330244,0.001806076,0.0003430385,0.0001404666,0.008329864,0.0001891362,0.000781836,0.1006687,0.2596908,0.09301044],"study_design_scores_gemma":[0.00703639,0.0005658244,0.2310526,0.008388776,0.00002899656,0.0001845875,0.002310281,0.03168023,0.00015654,0.3558359,0.3605351,0.002224648],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9326465,0.001693278,0.03483327,0.01112651,0.004242642,0.0007248711,0.00003051517,0.0001130361,0.01458944],"genre_scores_gemma":[0.9700003,0.00008099191,0.007363706,0.02089795,0.000398122,0.00003401345,0.00001022281,0.00004779621,0.00116689],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.3019717,"threshold_uncertainty_score":0.9999847,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.575652382298607,"score_gpt":0.5373958065628411,"score_spread":0.03825657573576591,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}