{"id":"W1910371136","doi":"10.21083/ajote.v1i1.1589","title":"Effective Test Administration in Schools: Principles and Good Practices for Test Administrators","year":2011,"lang":"en","type":"article","venue":"African Journal of Teacher Education","topic":"Education Systems and Policy","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Administration (probate law); Psychology; Construct (python library); Achievement test; Reliability (semiconductor); Construct validity; Academic achievement; Mathematics education; Medical education; Standardized test; Medicine; Political science; Psychometrics; Computer science; Clinical psychology","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001919071,0.00008582057,0.0001476999,0.0001754369,0.0001473134,0.00009312497,0.0001418554,0.00007529203,0.00003462191],"category_scores_gemma":[0.006983569,0.00007878483,0.00004117639,0.0002477435,0.00007769358,0.000564624,0.000006138326,0.0001785854,0.000003273744],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0001674326,"about_ca_system_score_gemma":0.002612037,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00112498,"about_ca_topic_score_gemma":0.001182219,"domain_scores_codex":[0.9989208,0.0002352865,0.0004005019,0.0001212875,0.0001689826,0.0001531583],"domain_scores_gemma":[0.9974319,0.0007402818,0.001253029,0.00008772036,0.0003138295,0.0001732574],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.00003089416,0.00111274,0.9076911,0.00002634188,0.00001228654,5.488794e-7,0.06836869,5.76437e-8,0.0001547136,0.01104037,0.0007164546,0.01084584],"study_design_scores_gemma":[0.0003084546,0.0008992721,0.5956784,0.0001247011,0.00005039839,0.00002780934,0.2948862,0.000001316454,0.0004264252,0.001479721,0.1059637,0.0001535092],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9672648,0.0003044267,0.00002640856,0.004543683,0.0004639601,0.0005311926,0.000004587689,0.000008789199,0.02685216],"genre_scores_gemma":[0.9959906,0.00003880745,0.001929168,0.0000455415,0.0009132733,0.00007375634,0.000001404823,0.000008876055,0.0009985811],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.3120126,"threshold_uncertainty_score":0.8360489,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.08482067924976232,"score_gpt":0.4063597155082588,"score_spread":0.3215390362584964,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}