{"id":"W4414254107","doi":"10.14445/23488387/ijcse-v12i8p104","title":"Importance of Structured Prompt Engineering to Generate Effective Test Cases","year":2025,"lang":"en","type":"article","venue":"International Journal of Computer Science and Engineering","topic":"Risk and Safety Analysis","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"Optech (Canada)","funders":"","keywords":"Task (project management); Test (biology); Software development process; Software; Software development; Test case; Applications of artificial intelligence","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01252757,0.0009727523,0.0003948998,0.002658509,0.0005846357,0.00272254,0.001423754,0.001113373,0.003316964],"category_scores_gemma":[0.1079188,0.000524912,0.0003787847,0.00084053,0.000931996,0.002749039,0.001295514,0.0013768,0.0008349666],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009534046,"about_ca_system_score_gemma":0.003296648,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001041334,"about_ca_topic_score_gemma":0.001930836,"domain_scores_codex":[0.978116,0.01350886,0.00140993,0.0009749359,0.005560677,0.0004295734],"domain_scores_gemma":[0.8497586,0.1170075,0.005555394,0.009280445,0.01722442,0.001173649],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0005588776,0.002572357,0.01868234,0.001172785,0.00006300912,0.0009632473,0.004067139,0.07696441,0.06087419,0.02093581,0.003677281,0.8094685],"study_design_scores_gemma":[0.0007514571,0.007980213,0.02659943,0.002035807,0.0002579805,0.003719477,0.004900069,0.6914069,0.143067,0.06955797,0.04943359,0.0002900089],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1675802,0.0003276896,0.8041297,0.001419446,0.0000926791,0.00352773,0.0002445162,0.004142459,0.01853565],"genre_scores_gemma":[0.3726911,0.0001731479,0.6245149,0.0002329691,0.00002436309,0.0006864575,0.0003154333,0.0003531168,0.001008411],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01252757,"threshold_uncertainty_score":0.06625289,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01335376537823231,"score_gpt":0.3105038821494064,"score_spread":0.2971501167711741,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}