{"id":"W4247890920","doi":"10.1109/metric.2004.1357925","title":"A prototype empirical evaluation of test driven development","year":2004,"lang":"en","type":"article","venue":"","topic":"Software Engineering Techniques and Practices","field":"Computer Science","cited_by":38,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Calgary","funders":"","keywords":"Test (biology); Computer science; Test-driven development; Programming language; Software development; Software","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.07897932,0.0007573161,0.0007347731,0.002036006,0.002003549,0.003161298,0.00383518,0.002610511,0.008800737],"category_scores_gemma":[0.3175451,0.0006986687,0.0007945684,0.002323817,0.003948974,0.0050468,0.003291865,0.002142872,0.001585856],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.004452237,"about_ca_system_score_gemma":0.004202288,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002332875,"about_ca_topic_score_gemma":0.001792415,"domain_scores_codex":[0.9067627,0.07449726,0.003698011,0.003168803,0.01023012,0.001643235],"domain_scores_gemma":[0.4402374,0.4592051,0.01375339,0.03250727,0.04900002,0.005296811],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.02721763,0.09525201,0.1398396,0.008857631,0.0006975549,0.001930866,0.05550411,0.01850957,0.01347739,0.07336676,0.02543502,0.5399119],"study_design_scores_gemma":[0.03596307,0.2607215,0.2776798,0.004988925,0.001218276,0.001851231,0.05114776,0.08404727,0.02474855,0.04931795,0.2074033,0.0009123248],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8955066,0.0006573652,0.03555509,0.002147753,0.0003144679,0.01787651,0.001041218,0.0003736136,0.04652733],"genre_scores_gemma":[0.9316412,0.000355693,0.04892553,0.0008305661,0.000089933,0.01464047,0.0007214084,0.00007635755,0.002718893],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.07897932,"threshold_uncertainty_score":0.4176875,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07390805756743042,"score_gpt":0.3532777063614152,"score_spread":0.2793696487939848,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}