{"id":"W4408054097","doi":"10.1080/02602938.2025.2468848","title":"Enhancing the structure of feedback forms increases trustworthiness and usefulness of peer feedback","year":2025,"lang":"en","type":"article","venue":"Assessment & Evaluation in Higher Education","topic":"Student Assessment and Feedback","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of British Columbia","funders":"","keywords":"Peer feedback; Trustworthiness; Psychology; Peer evaluation; Peer review; Higher education; Negative feedback; Social psychology; Mathematics education; Political science; Engineering","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.04074021,0.0008198669,0.0009751558,0.001875741,0.001147223,0.002984375,0.0009414675,0.00102671,0.00291411],"category_scores_gemma":[0.2957738,0.0006086372,0.0008304684,0.0008412204,0.00113804,0.003032153,0.0025681,0.00123618,0.0007580438],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001070927,"about_ca_system_score_gemma":0.002729222,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0007184987,"about_ca_topic_score_gemma":0.001082351,"domain_scores_codex":[0.9261714,0.04651383,0.006238387,0.002983245,0.0169679,0.001125259],"domain_scores_gemma":[0.5634441,0.3327027,0.03281469,0.03055443,0.03578961,0.004694568],"domain_codex":null,"domain_gemma":"evaluation","domain_candidate":"evaluation","domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.002857209,0.004328356,0.09502631,0.003014153,0.0002858604,0.0003301967,0.03168365,0.002519313,0.05305385,0.001516421,0.002378114,0.8030065],"study_design_scores_gemma":[0.002833156,0.04728124,0.7076008,0.006667865,0.001834908,0.002517334,0.02390664,0.03337772,0.1091016,0.01259832,0.05133035,0.0009500515],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9248818,0.000551374,0.05742595,0.0008828653,0.0001501832,0.003604962,0.0001022047,0.001038636,0.01136191],"genre_scores_gemma":[0.8911409,0.0003940963,0.1054995,0.0001812913,0.0001291895,0.001307448,0.00007272259,0.0001302371,0.0011445],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.9592598,"threshold_uncertainty_score":0.2154574,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03977582508428822,"score_gpt":0.4023611073920225,"score_spread":0.3625852823077343,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}