{"id":"W4391036726","doi":"10.4236/ojepi.2024.141005","title":"Cautionary Remarks When Testing Agreement between Two Raters for Continuous Scale Measurements: A Tutorial in Clinical Epidemiology with Implementation Using R","year":2024,"lang":"en","type":"article","venue":"Open Journal of Epidemiology","topic":"Reliability and Agreement in Measurement","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Western University","funders":"","keywords":"Statistic; Scale (ratio); Sample size determination; Statistics; Sample (material); Linear regression; Degrees of freedom (physics and chemistry); Data mining; Computer science; Medicine; Econometrics; Mathematics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2590153,0.004684025,0.003911904,0.007139266,0.002613325,0.006831243,0.01033028,0.01314746,0.01763214],"category_scores_gemma":[0.7338336,0.003457085,0.005454313,0.00715623,0.01340868,0.01017652,0.006337423,0.02868109,0.0241778],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003673046,"about_ca_system_score_gemma":0.007006316,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003798869,"about_ca_topic_score_gemma":0.004508532,"domain_scores_codex":[0.5720521,0.3391851,0.03774661,0.01193051,0.03746323,0.001622436],"domain_scores_gemma":[0.1643767,0.7314397,0.0187756,0.0320491,0.05117281,0.002186094],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0003696117,0.000124304,0.001782716,0.004738502,0.00044382,0.001005898,0.00396811,0.002092232,0.001243955,0.0294509,0.8156708,0.1391091],"study_design_scores_gemma":[0.0003893705,0.0003508952,0.003238166,0.01578132,0.0004467498,0.004215175,0.001483273,0.01670123,0.00543749,0.1403502,0.8109173,0.0006887845],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.001827203,0.01329042,0.6761305,0.228365,0.04817324,0.001286719,0.001904255,0.01981892,0.009203664],"genre_scores_gemma":[0.01777977,0.005674726,0.8270408,0.1068915,0.01399518,0.00614312,0.0005630176,0.01012034,0.01179152],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.7409847,"threshold_uncertainty_score":0.9137661,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.7023443814366066,"score_gpt":0.5922363059400028,"score_spread":0.1101080754966038,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}