{"id":"W2808880545","doi":"10.15760/nwjte.2011.9.2.7","title":"Towards Balanced Assessment of Student Teaching Performance","year":2011,"lang":"en","type":"article","venue":"Northwest Journal of Teacher Education","topic":"Teacher Education and Leadership Studies","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Summative assessment; Formative assessment; Context (archaeology); Knowledge survey; Process (computing); Medical education; Pedagogy; Mathematics education; Psychology; Computer science; Medicine","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0545303,0.0009754225,0.001000173,0.004699567,0.001485804,0.007935554,0.002020304,0.0015838,0.001964022],"category_scores_gemma":[0.09586094,0.0006267686,0.0004337247,0.002896403,0.001986406,0.006032747,0.008560143,0.003197181,0.001482457],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003221997,"about_ca_system_score_gemma":0.007090971,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003705953,"about_ca_topic_score_gemma":0.00639952,"domain_scores_codex":[0.9328169,0.02916091,0.006103198,0.00452248,0.0262591,0.001137352],"domain_scores_gemma":[0.9316273,0.01253414,0.007445246,0.005068344,0.04126463,0.002060296],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0005210923,0.0007089623,0.08166955,0.0007685907,0.0001023983,0.0001389501,0.01830171,0.002830726,0.02130729,0.01208272,0.007637798,0.8539303],"study_design_scores_gemma":[0.0003931971,0.006207881,0.4832219,0.004945427,0.000345772,0.00175361,0.04121862,0.04994148,0.1049466,0.09263509,0.2135101,0.0008802389],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3815095,0.001960447,0.5162272,0.0121917,0.0008461474,0.00259387,0.0008355449,0.003468119,0.08036753],"genre_scores_gemma":[0.4610326,0.000928299,0.5257799,0.000764046,0.0001118315,0.001455876,0.000622866,0.0003593829,0.008945164],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.0545303,"threshold_uncertainty_score":0.2883873,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.124465235721844,"score_gpt":0.4340206400967623,"score_spread":0.3095554043749183,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}