{"id":"W2318026979","doi":"10.15405/futureacademy/ejsbs(2301-2218).2012.2.13","title":"Why Consistency Is Not Possible In Experienced Teacher Evaluations","year":2012,"lang":"en","type":"article","venue":"The European Journal of Social & Behavioural Sciences","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Psychology; Task (project management); Principal (computer security); Performance appraisal; Process (computing); Christian ministry; Phenomenon; Grounded theory; Employee Performance Appraisal; Pedagogy; Applied psychology; Social psychology; Mathematics education; Computer science; Management; Qualitative research; Epistemology; Political science; Medicine; Nursing","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.03387627,0.0001255917,0.0002286486,0.0002571628,0.001060451,0.0003696665,0.001351643,0.00002366167,0.002915494],"category_scores_gemma":[0.0007025682,0.00006673868,0.0001712476,0.001210463,0.0007603762,0.001485012,0.0001202416,0.0002907983,0.0001829797],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00007733759,"about_ca_system_score_gemma":0.0002695987,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00004272254,"about_ca_topic_score_gemma":0.00003016264,"domain_scores_codex":[0.9933594,0.00223127,0.00102904,0.0001688492,0.002829066,0.0003823526],"domain_scores_gemma":[0.9981518,0.0002994677,0.0007888203,0.0001742386,0.000449075,0.0001365336],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.00005395821,0.0004106948,0.6965452,0.000001330713,0.00001521888,0.00001098757,0.1376254,0.0001294984,0.003114007,0.00234253,0.02847274,0.1312784],"study_design_scores_gemma":[0.0004321387,0.0001591611,0.964356,0.00001481121,0.00002461835,0.00003049542,0.02854084,0.0001749354,0.0003926586,0.0006045068,0.005138783,0.0001310338],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9539205,0.0001611803,0.0001192219,0.0149913,0.000734977,0.0001120257,0.000002606624,0.000006701776,0.0299515],"genre_scores_gemma":[0.9966423,0.00001093857,0.0003484591,0.001785289,0.000349156,0.000002192702,2.447887e-7,0.000006256029,0.0008551667],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.2678108,"threshold_uncertainty_score":0.997996,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3937768503371043,"score_gpt":0.5170564721557023,"score_spread":0.123279621818598,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}