{"id":"W4403361536","doi":"10.1016/j.healun.2024.10.007","title":"Diagnostic alignment to optimize inter-rater reliability among lung transplant pathologists","year":2024,"lang":"en","type":"article","venue":"The Journal of Heart and Lung Transplantation","topic":"Transplantation: Methods and Outcomes","field":"Medicine","cited_by":2,"is_retracted":false,"has_abstract":false,"ca_institutions":"Toronto General Hospital; University Health Network","funders":"National Institute of Allergy and Infectious Diseases; Cystic Fibrosis Foundation","keywords":"Medicine; Histopathology; Lung transplantation; Bronchiolitis; Lung; Inter-rater reliability; Confidence interval; Internal medicine; Pathology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.09562331,0.0009803361,0.001275136,0.002678795,0.002957254,0.003937101,0.002473803,0.001308084,0.004130183],"category_scores_gemma":[0.2470422,0.001194718,0.001331339,0.003762912,0.001234685,0.002957146,0.003319876,0.001525151,0.002090073],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001361001,"about_ca_system_score_gemma":0.006083289,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002701284,"about_ca_topic_score_gemma":0.007972023,"domain_scores_codex":[0.8753526,0.08695479,0.01432298,0.01207742,0.009294763,0.001997458],"domain_scores_gemma":[0.791234,0.09195425,0.02041752,0.03310827,0.06186197,0.001424074],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.004474109,0.001043122,0.3715183,0.00199771,0.002200789,0.0004437179,0.01917959,0.008043905,0.04045692,0.01152336,0.02212106,0.5169974],"study_design_scores_gemma":[0.001699534,0.003357463,0.4933858,0.001500411,0.004783495,0.002418048,0.01330714,0.231687,0.1302553,0.03581769,0.08116315,0.000624915],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.2961846,0.001306956,0.6819602,0.001353585,0.0008678451,0.003805994,0.00150574,0.003518843,0.009496202],"genre_scores_gemma":[0.5502655,0.0001851051,0.4417741,0.0004730752,0.0002407433,0.003305677,0.001593712,0.0007118739,0.001450166],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.9043767,"threshold_uncertainty_score":0.5057104,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01322369181395695,"score_gpt":0.3135123341048449,"score_spread":0.300288642290888,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}