{"id":"W7116754798","doi":"10.1177/10731911251391563","title":"Embedded Validity Scales to Examine Caregiver Response Styles When Measuring Infant/Toddler Developmental Status","year":2025,"lang":"en","type":"article","venue":"Assessment","topic":"Infant Development and Preterm Care","field":"Medicine","cited_by":1,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Windsor","funders":"Eunice Kennedy Shriver National Institute of Child Health and Human Development","keywords":"Response bias; Test validity; Reliability (semiconductor); Validity; Psychometrics; Concurrent validity; Ethnic group; Scale (ratio)","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.0007941126,0.0003095907,0.0004394048,0.0003756731,0.0001946913,0.00007851121,0.0001733863,0.0001214389,0.0002902005],"category_scores_gemma":[0.0001995933,0.0002820759,0.0001001513,0.0003065426,0.00005248935,0.0001493431,0.00027094,0.0002526074,0.00003147269],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009153722,"about_ca_system_score_gemma":0.001028425,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00007354,"about_ca_topic_score_gemma":0.0001040148,"domain_scores_codex":[0.9976736,0.0002021934,0.0004411462,0.0004998525,0.0006204922,0.0005627655],"domain_scores_gemma":[0.9988022,0.0002346285,0.00007289902,0.000374639,0.0002548326,0.0002607442],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.004095724,0.0001764289,0.8613002,0.0004608884,0.0008064584,0.0004117306,0.02144351,0.00001546906,0.04605404,0.0006666878,0.04081088,0.02375794],"study_design_scores_gemma":[0.002249897,0.0001888033,0.8752747,0.0005091494,0.0001011393,0.00002231175,0.002904017,0.00003157056,0.01809482,0.00007469668,0.1000937,0.0004552056],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9570425,0.0000826898,0.0004654313,0.0009224655,0.0003769762,0.0007986191,0.00005746114,0.000142696,0.04011122],"genre_scores_gemma":[0.9621596,0.00002581382,0.03285186,0.001267794,0.00006278586,0.0001276851,0.0001825038,0.00002706008,0.00329488],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.05928285,"threshold_uncertainty_score":0.9999632,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04941604682889743,"score_gpt":0.3241424135637946,"score_spread":0.2747263667348971,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}