{"id":"W7132913430","doi":"","title":"Conflicts and challenges of testing: Differences between novice and experienced teachers&apos; preparation practices for large-scale assessments","year":2008,"lang":"","type":"dissertation","venue":"TSpace","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Set (abstract data type); Test (biology); Sample (material); Educational measurement","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.001075691,0.0005294115,0.0009596639,0.0001639293,0.001130845,0.000263785,0.0003976107,0.0005856809,0.00004253576],"category_scores_gemma":[0.0009186757,0.0005290589,0.00009454961,0.0003324885,0.0003819033,0.0008008837,0.000102166,0.000303702,0.000002118592],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00005564183,"about_ca_system_score_gemma":0.0003422821,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001714782,"about_ca_topic_score_gemma":0.002534452,"domain_scores_codex":[0.996383,0.0003633044,0.0006768872,0.0009941987,0.0009818685,0.0006007256],"domain_scores_gemma":[0.993988,0.002186265,0.00276482,0.0002498961,0.0005730139,0.0002380574],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"qualitative","study_design_gemma":"qualitative","study_design_scores_codex":[0.0001981967,0.0003031237,0.1464459,0.0005161372,0.0002322806,8.554089e-7,0.844487,3.249304e-7,0.001593056,0.0005600672,0.00003791276,0.005625156],"study_design_scores_gemma":[0.00128681,0.001350947,0.2932833,0.0003256407,0.0004027492,9.930476e-7,0.6973668,0.0003322482,0.0004425628,0.00006732871,0.004575304,0.0005652519],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.971254,0.005034452,0.00008262192,0.0005218019,0.0003318843,0.001848044,0.00002028958,0.00004312687,0.02086377],"genre_scores_gemma":[0.9766228,0.01358047,0.002505868,0.00001439103,0.0003570569,0.0002382815,0.00009970773,0.00003456789,0.006546832],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.1471201,"threshold_uncertainty_score":0.9997161,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2275064307269525,"score_gpt":0.5093735031714349,"score_spread":0.2818670724444824,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}