{"id":"W2255858454","doi":"10.1016/j.envint.2016.01.005","title":"How credible are the study results? Evaluating and applying internal validity tools to literature-based assessments of environmental health hazards","year":2016,"lang":"en","type":"article","venue":"Environment International","topic":"Health Systems, Economic Evaluations, Quality of Life","field":"Economics, Econometrics and Finance","cited_by":92,"is_retracted":false,"has_abstract":false,"ca_institutions":"McMaster University","funders":"Agency for Healthcare Research and Quality; National Institutes of Health; U.S. Environmental Protection Agency","keywords":"Credibility; Risk assessment; Grading (engineering); Environmental health; Observational study; Risk analysis (engineering); Hazard; Exposure assessment; Agency (philosophy); Internal validity; Public health; Applied psychology; Psychology; Computer science; Medicine; Engineering; Political science","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.6462011,0.001908007,0.003945795,0.01580052,0.003679879,0.01221799,0.00491997,0.004264208,0.001592581],"category_scores_gemma":[0.8624983,0.001862157,0.008383957,0.009901274,0.01860778,0.0126407,0.01081473,0.00510919,0.0002678283],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00683136,"about_ca_system_score_gemma":0.01021383,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002950688,"about_ca_topic_score_gemma":0.003688921,"domain_scores_codex":[0.3179444,0.5567229,0.05462909,0.01170458,0.05633944,0.002659536],"domain_scores_gemma":[0.02823037,0.8851902,0.03943465,0.02163803,0.0245318,0.0009749008],"domain_codex":"methods","domain_gemma":"methods","domain_candidate":"methods","domain_consensus":"methods","study_design_codex":"observational","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00238362,0.001199154,0.6981245,0.004915421,0.01998601,0.0003593939,0.04592876,0.01108383,0.0008499475,0.05100282,0.004442359,0.1597242],"study_design_scores_gemma":[0.002595958,0.005561718,0.4917008,0.01208057,0.01407849,0.0009873876,0.03714636,0.1321628,0.009939933,0.2757122,0.01668123,0.001352548],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6892449,0.008934439,0.2454823,0.01299912,0.001495347,0.007306873,0.001061928,0.0003491021,0.03312606],"genre_scores_gemma":[0.9492155,0.0005559478,0.04537117,0.0007307657,0.00022895,0.003328336,0.0002959213,0.00006600693,0.000207444],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.3537989,"threshold_uncertainty_score":0.4362971,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3790899069598085,"score_gpt":0.445247100886897,"score_spread":0.06615719392708841,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}