{"id":"W7132913430","doi":"","title":"Conflicts and challenges of testing: Differences between novice and experienced teachers&apos; preparation practices for large-scale assessments","year":2008,"lang":"","type":"dissertation","venue":"TSpace","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Set (abstract data type); Test (biology); Sample (material); Educational measurement","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004388678,0.0001461129,0.0002240901,0.0006672579,0.002157069,0.002240354,0.0007029501,0.0004632999,0.001173854],"category_scores_gemma":[0.03551702,0.0003236962,0.0001691843,0.0006057356,0.001928165,0.0007290787,0.001643164,0.000694947,0.0001958488],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.004211023,"about_ca_system_score_gemma":0.006520617,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.1940831,"about_ca_topic_score_gemma":0.39626,"domain_scores_codex":[0.9954021,0.001450415,0.0004025868,0.0002658309,0.001705511,0.0007734774],"domain_scores_gemma":[0.9720408,0.01042603,0.007362336,0.0008698499,0.00465217,0.004648837],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"qualitative","study_design_scores_codex":[0.0001342675,0.0001660779,0.714967,0.00005134055,0.00001268108,0.0005609963,0.2508354,0.000102754,0.001794214,0.0001587385,0.0006811164,0.03053549],"study_design_scores_gemma":[0.000008335101,0.0001419282,0.827725,0.0000431892,0.000005396409,0.0002351387,0.1691379,0.0002194151,0.0003083853,0.0001279546,0.002026986,0.0000202642],"study_design_candidate":"qualitative","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9988924,0.00004353474,0.00006899673,0.0001300482,0.000002860661,0.000009409707,0.00001027295,0.00000292505,0.0008395807],"genre_scores_gemma":[0.9993848,0.00003760574,0.0001108882,0.00003365021,0.000001885514,0.00000644744,0.00001348548,0.000001966219,0.0004092799],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.1940831,"threshold_uncertainty_score":0.3859069,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2275064307269525,"score_gpt":0.5093735031714349,"score_spread":0.2818670724444824,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}