{"id":"W4256693308","doi":"10.15760/nwjte.2012.9.2.7","title":"Towards Balanced Assessment of Student Teaching Performance","year":2012,"lang":"en","type":"article","venue":"Northwest Journal of Teacher Education","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Lethbridge","funders":"","keywords":"Summative assessment; Formative assessment; Context (archaeology); Knowledge survey; Process (computing); Medical education; Pedagogy; Mathematics education; Psychology; Computer science; Medicine","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002556569,0.00009835221,0.0002092169,0.0001349105,0.0002091835,0.00004476038,0.0003344094,0.00006068978,0.0002134656],"category_scores_gemma":[0.00005706799,0.00008503296,0.0001009965,0.0001481003,0.00006659231,0.0007926704,0.00002986563,0.0003291148,0.000004446483],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0003458275,"about_ca_system_score_gemma":0.0009299834,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0004240842,"about_ca_topic_score_gemma":0.00008760513,"domain_scores_codex":[0.9981802,0.0002444742,0.0004494546,0.00007308231,0.0008047936,0.000248071],"domain_scores_gemma":[0.9987763,0.00004412356,0.0006359399,0.0001176352,0.0002722231,0.0001537825],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.000005489669,0.001120967,0.9420812,0.00001075473,0.00003422633,9.675838e-8,0.02481431,0.000004997924,0.0001978285,0.002096053,0.0004791879,0.02915491],"study_design_scores_gemma":[0.0002176536,0.00008441656,0.9658754,0.00005409967,0.00005611499,0.000002520689,0.020276,0.000003131552,0.00003938329,0.00002085672,0.01328306,0.00008735428],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9506296,0.0003070009,0.0001113399,0.0005127673,0.001973891,0.0001262539,5.37791e-7,0.000008767931,0.04632982],"genre_scores_gemma":[0.9936004,0.0001771519,0.003821383,0.00004657768,0.00163972,0.000005679708,0.000003427744,0.00000896188,0.0006966431],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.04563318,"threshold_uncertainty_score":0.3467543,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0344671714649823,"score_gpt":0.4169282447697099,"score_spread":0.3824610733047276,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}