{"id":"W2097681968","doi":"10.1093/biostatistics/2.3.323","title":"Efficiency considerations in the analysis of inter-observer agreement","year":2001,"lang":"en","type":"article","venue":"Biostatistics","topic":"Reliability and Agreement in Measurement","field":"Decision Sciences","cited_by":18,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Guelph","funders":"","keywords":"Cohen's kappa; Replicate; Statistics; Inter-rater reliability; Reliability (semiconductor); Statistic; Kappa; Mathematics; Binary number; Computer science; Binomial distribution; Econometrics; Arithmetic; Physics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.3907782,0.002745539,0.00455962,0.004122018,0.001741974,0.005607869,0.005316505,0.004008732,0.002050049],"category_scores_gemma":[0.6833754,0.001955993,0.004840995,0.005029655,0.00791999,0.008073524,0.007114369,0.005776657,0.00133425],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.004075131,"about_ca_system_score_gemma":0.003366545,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00259779,"about_ca_topic_score_gemma":0.001517115,"domain_scores_codex":[0.4759949,0.4531517,0.01780426,0.02028565,0.03061931,0.002144159],"domain_scores_gemma":[0.134743,0.8021755,0.01427239,0.03551462,0.0128765,0.0004179615],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.001838594,0.0004143204,0.04883049,0.002479892,0.006464876,0.001113154,0.006440538,0.158596,0.0049408,0.460951,0.006433282,0.3014971],"study_design_scores_gemma":[0.000221427,0.001067805,0.02775201,0.0009114322,0.0008723466,0.001304071,0.0009919952,0.4585225,0.007953979,0.4808466,0.01923726,0.0003185856],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.004955758,0.0006868019,0.9924064,0.0003754382,0.00007683824,0.0003040353,0.0000650285,0.0001176913,0.00101188],"genre_scores_gemma":[0.3186735,0.0008259462,0.6725587,0.001057141,0.0003105336,0.002962477,0.0004508146,0.0007040234,0.002456878],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.6092218,"threshold_uncertainty_score":0.7512789,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2311350697927032,"score_gpt":0.4052181703110456,"score_spread":0.1740831005183424,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}