{"id":"W4238524924","doi":"10.26434/chemrxiv-2021-7m4tw","title":"Response process validity evidence in chemistry education research","year":2021,"lang":"en","type":"preprint","venue":"ChemRxiv","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Ottawa","funders":"","keywords":"Process (computing); Cognition; Psychology; Management science; Applied psychology; Computer science; Engineering; Neuroscience","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.7174793,0.001932525,0.00318522,0.01068812,0.006126037,0.01729918,0.006584293,0.0101565,0.01877145],"category_scores_gemma":[0.9097313,0.002678048,0.006634715,0.01343722,0.02337153,0.01856312,0.0128464,0.01074715,0.003869088],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.01215295,"about_ca_system_score_gemma":0.02424206,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004120864,"about_ca_topic_score_gemma":0.00351375,"domain_scores_codex":[0.1624894,0.5808335,0.08518219,0.02846259,0.1387773,0.004255024],"domain_scores_gemma":[0.02505879,0.8499393,0.02809495,0.04692259,0.0490756,0.000908732],"domain_codex":"methods","domain_gemma":"methods","domain_candidate":"methods","domain_consensus":"methods","study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.002250948,0.0009223929,0.07374008,0.06403147,0.005087323,0.0004696216,0.04232265,0.003154399,0.001874363,0.3507641,0.02188767,0.433495],"study_design_scores_gemma":[0.001613955,0.002760187,0.07121418,0.1425077,0.003963864,0.001229807,0.01901607,0.0108506,0.01376545,0.4633859,0.2689931,0.0006993184],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.07328189,0.07617807,0.5668975,0.09664524,0.008472991,0.01575541,0.003750926,0.0007744218,0.1582435],"genre_scores_gemma":[0.6447381,0.0198059,0.2677731,0.02667791,0.002456556,0.02942166,0.003540013,0.001486122,0.004100735],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.2825207,"threshold_uncertainty_score":0.3483983,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.7160760744590771,"score_gpt":0.6518589554735377,"score_spread":0.06421711898553939,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}