{"id":"W2477109771","doi":"10.4018/978-1-4666-4999-6.ch020","title":"Reliability and Validity in Automated Content Analysis","year":2014,"lang":"en","type":"book-chapter","venue":"Advances in linguistics and communication studies","topic":"Computational and Text Analysis Methods","field":"Social Sciences","cited_by":8,"is_retracted":false,"has_abstract":true,"ca_institutions":"McGill University","funders":"","keywords":"Diction; Reliability (semiconductor); Content (measure theory); Content validity; Volume (thermodynamics); Computer science; Validity; Data science; Statistics; Linguistics; Mathematics; Philosophy; Psychometrics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.2237551,0.001452949,0.002233883,0.02132576,0.004417751,0.01709754,0.003558009,0.002710432,0.004926932],"category_scores_gemma":[0.5236157,0.001916668,0.001944923,0.01778344,0.02429454,0.02181761,0.01142912,0.005048086,0.002677672],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.007900943,"about_ca_system_score_gemma":0.00812961,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005002895,"about_ca_topic_score_gemma":0.00305682,"domain_scores_codex":[0.6708732,0.2338313,0.01572125,0.01625695,0.06122718,0.002090076],"domain_scores_gemma":[0.3026427,0.5702187,0.01495731,0.05397999,0.05723653,0.0009647055],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001274021,0.00009940448,0.01457863,0.00209395,0.0003763266,0.0001538406,0.03155476,0.003756921,0.0008094491,0.5953022,0.01284173,0.3383054],"study_design_scores_gemma":[0.00005405695,0.0001069751,0.01423799,0.003610732,0.0001363409,0.0003786415,0.007243209,0.0178704,0.001865722,0.8889499,0.06535082,0.0001952358],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.04209846,0.01940383,0.73768,0.01346469,0.001500359,0.002559615,0.0008077083,0.001535977,0.1809494],"genre_scores_gemma":[0.4831008,0.007330236,0.4900938,0.002128566,0.001680358,0.005077622,0.001123886,0.001282342,0.008182294],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.2237551,"threshold_uncertainty_score":0.9572482,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1640921458624373,"score_gpt":0.449471534461814,"score_spread":0.2853793885993767,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}