{"id":"W767205952","doi":"10.26522/tl.v1i2.103","title":"Developing a Diagnostic Assessment Model","year":2003,"lang":"en","type":"article","venue":"Teaching and Learning","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Mathematics education; Standardized test; Quality (philosophy); Scale (ratio); Educational assessment; Test (biology); Psychology; Medical education; Medicine; Geography","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.007664173,0.001019813,0.0006922111,0.003076563,0.001035301,0.003702552,0.002926722,0.002185274,0.009266445],"category_scores_gemma":[0.03135784,0.000548382,0.001002204,0.001529839,0.001484481,0.004536485,0.002445767,0.002051954,0.002737399],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002952904,"about_ca_system_score_gemma":0.003773158,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01028234,"about_ca_topic_score_gemma":0.005980832,"domain_scores_codex":[0.9954824,0.002316315,0.0003136619,0.0007066469,0.0009063697,0.0002744969],"domain_scores_gemma":[0.9886039,0.006581801,0.0005368411,0.0005905137,0.00334727,0.0003396218],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002156064,0.0002764478,0.01623951,0.0002861812,0.0001694663,0.0004967834,0.0009997678,0.1663742,0.0008491145,0.4376983,0.01106225,0.3653325],"study_design_scores_gemma":[0.00004800608,0.00008033904,0.0008504776,0.00009755359,0.00005093358,0.0001876926,0.0002564536,0.7412632,0.0004473821,0.2469475,0.009741814,0.00002870415],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01087886,0.0002600231,0.971442,0.001863852,0.00007551786,0.0003168387,0.0004757488,0.0006412416,0.01404589],"genre_scores_gemma":[0.3721277,0.0003982835,0.6194655,0.0004007849,0.00007849611,0.000858242,0.001214503,0.0000725681,0.005384078],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01028234,"threshold_uncertainty_score":0.04053253,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1991485901968059,"score_gpt":0.4974074992217229,"score_spread":0.2982589090249169,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}