{"id":"W2594171765","doi":"10.1016/j.annemergmed.2016.12.025","title":"Individual Gestalt Is Unreliable for the Evaluation of Quality in Medical Education Blogs: A METRIQ Study","year":2017,"lang":"en","type":"article","venue":"Annals of Emergency Medicine","topic":"Child Therapy and Development","field":"Psychology","cited_by":51,"is_retracted":false,"has_abstract":false,"ca_institutions":"McMaster University; University of Toronto; Western University; University of Alberta; University of Saskatchewan","funders":"","keywords":"Generalizability theory; Intraclass correlation; Gestalt psychology; Medicine; Likert scale; Pearson product-moment correlation coefficient; Reliability (semiconductor); Correlation; Family medicine; Clinical psychology; Statistics; Psychometrics; Psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.03165615,0.0002909687,0.0006767177,0.002981568,0.001548335,0.003925601,0.00126092,0.001754477,0.004673311],"category_scores_gemma":[0.2190778,0.0004685194,0.0008470321,0.002024557,0.002648339,0.005163365,0.003396582,0.0019613,0.001328619],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001715534,"about_ca_system_score_gemma":0.0008354181,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00331871,"about_ca_topic_score_gemma":0.003739507,"domain_scores_codex":[0.9775264,0.01260902,0.002121287,0.001155025,0.005550073,0.001038075],"domain_scores_gemma":[0.5808452,0.294077,0.06084816,0.02442767,0.03260251,0.007199336],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0008069034,0.0004850841,0.9702346,0.00009701771,0.0001908541,0.00006400778,0.01200614,0.00009654312,0.0002596405,0.0005512282,0.0009168215,0.01429114],"study_design_scores_gemma":[0.00007151043,0.0005054896,0.9821984,0.00008768361,0.0001413157,0.0001752809,0.01254251,0.0009725943,0.0003484481,0.0008380549,0.002068433,0.00005025177],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9932955,0.0001104499,0.0005540093,0.0005247093,0.0000270792,0.00008079516,0.0002832319,0.00001576372,0.005108331],"genre_scores_gemma":[0.9987141,0.00002865907,0.0002838605,0.0001848423,0.00002753153,0.00006430306,0.0001243728,0.00003089042,0.0005415088],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.9683439,"threshold_uncertainty_score":0.1674157,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.4755115456714767,"score_gpt":0.5955290715001829,"score_spread":0.1200175258287062,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}