{"id":"W4366384499","doi":"10.3138/cjpe.0025.013","title":"Learning from Evaluation Misadventures: The Importance of Good Communication","year":2011,"lang":"en","type":"article","venue":"Canadian Journal of Program Evaluation","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Variety (cybernetics); Relevance (law); Key (lock); Order (exchange); Quality (philosophy); Government (linguistics); Knowledge management; Process management; Management science; Computer science; Business; Psychology; Risk analysis (engineering); Engineering ethics; Political science; Engineering; Epistemology; Artificial intelligence; Computer security","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.03068425,0.0001131966,0.000223231,0.0003335026,0.0003004423,0.0001465774,0.0008719491,0.00008327534,0.006040428],"category_scores_gemma":[0.005285184,0.00007468922,0.0001461994,0.0006598761,0.0001336039,0.0006088484,0.00002029397,0.0003222617,0.00003645787],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0002780368,"about_ca_system_score_gemma":0.003590181,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001690682,"about_ca_topic_score_gemma":0.03222033,"domain_scores_codex":[0.9934087,0.0020265,0.001260316,0.0001719098,0.002941518,0.0001910768],"domain_scores_gemma":[0.9918113,0.0004937689,0.001926124,0.0005817923,0.004982213,0.0002047549],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.00002768504,0.00005768621,0.2747269,0.000002246118,0.00006027301,7.164481e-7,0.01128495,0.003020111,0.00005550102,0.0007729086,0.0009532202,0.7090378],"study_design_scores_gemma":[0.001404284,0.000640767,0.7661749,0.0001024008,0.0003700731,0.00001325368,0.01257116,0.1340512,0.0004847384,0.06894502,0.01506785,0.0001744416],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9834912,0.002564693,0.0005850393,0.001264219,0.0005171812,0.001314806,0.000005476879,0.000007212914,0.01025012],"genre_scores_gemma":[0.9965019,0.00003682312,0.003092759,0.0001116628,0.00008290409,0.00007113209,0.00003362245,0.000008825474,0.00006036025],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.7088634,"threshold_uncertainty_score":0.9981145,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.4890784253463381,"score_gpt":0.5126845869829371,"score_spread":0.02360616163659901,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}