{"id":"W2318026979","doi":"10.15405/futureacademy/ejsbs(2301-2218).2012.2.13","title":"Why Consistency Is Not Possible In Experienced Teacher Evaluations","year":2012,"lang":"en","type":"article","venue":"The European Journal of Social & Behavioural Sciences","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Consistency (knowledge bases); Psychology; Task (project management); Principal (computer security); Performance appraisal; Process (computing); Christian ministry; Phenomenon; Grounded theory; Employee Performance Appraisal; Pedagogy; Applied psychology; Social psychology; Mathematics education; Computer science; Management; Qualitative research; Epistemology; Political science; Medicine; Nursing","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.3092616,0.00087907,0.002317202,0.004599632,0.003878286,0.01050656,0.004698615,0.003611155,0.002430386],"category_scores_gemma":[0.6995941,0.001641225,0.001248603,0.00331789,0.0096146,0.01230389,0.008008682,0.005389981,0.0007042405],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.008008338,"about_ca_system_score_gemma":0.00708043,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.006150278,"about_ca_topic_score_gemma":0.004268539,"domain_scores_codex":[0.4634955,0.3339906,0.05274307,0.02876032,0.1134676,0.007542946],"domain_scores_gemma":[0.2614932,0.4727361,0.06305351,0.06523217,0.131915,0.005569968],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.002226568,0.0007494607,0.2072798,0.006355224,0.002058454,0.0009717003,0.1186876,0.005432805,0.00445515,0.1175228,0.04726563,0.4869949],"study_design_scores_gemma":[0.001129701,0.002692793,0.2812625,0.01180516,0.0007574112,0.002399224,0.05638814,0.02212902,0.01205581,0.5159787,0.0922029,0.001198679],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.5441941,0.01632398,0.2332038,0.1021791,0.004336624,0.003280084,0.0008956615,0.001736518,0.09385017],"genre_scores_gemma":[0.9615669,0.0006806292,0.03077158,0.003541318,0.0003918606,0.001310311,0.0002223685,0.0002194969,0.001295399],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.3092616,"threshold_uncertainty_score":0.8518034,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3937768503371043,"score_gpt":0.5170564721557023,"score_spread":0.123279621818598,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}