{"id":"W4406603690","doi":"10.48550/arxiv.2501.09888","title":"Understanding the Effectiveness of LLMs in Automated Self-Admitted Technical Debt Repayment","year":2025,"lang":"en","type":"preprint","venue":"ArXiv.org","topic":"FinTech, Crowdfunding, Digital Finance","field":"Business, Management and Accounting","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Debt; Economics; Business; Actuarial science; Finance","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.001753114,0.0004816596,0.0007101945,0.0005971653,0.0001402606,0.0002095094,0.0009500901,0.0004216748,0.00002096479],"category_scores_gemma":[0.000729056,0.0004008237,0.0002481873,0.001395042,0.0001636248,0.0004146713,0.00219474,0.0008586158,0.00006187894],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007756632,"about_ca_system_score_gemma":0.0001455062,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0005904683,"about_ca_topic_score_gemma":0.0001836591,"domain_scores_codex":[0.9974395,0.00009047561,0.0007755008,0.0008194027,0.0003796864,0.0004954055],"domain_scores_gemma":[0.9978228,0.000525189,0.0006032903,0.0008768806,0.0001586842,0.00001321936],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.0002325547,0.0003811909,0.9421108,0.003962915,0.0001715063,0.00005412865,0.00004702119,0.0009087202,0.001495603,0.04942933,0.001084605,0.000121567],"study_design_scores_gemma":[0.0009818531,0.00002567328,0.9621919,0.005179426,0.000177795,0.000003403899,0.0001178808,0.008872237,0.001547375,0.01863589,0.001552164,0.0007143453],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9617565,0.0001524825,0.00223873,0.0003147751,0.0008969674,0.001521552,0.00001327507,0.001386703,0.03171903],"genre_scores_gemma":[0.9991128,0.00002540735,0.00009931309,0.0002066048,0.0001245437,0.0001991139,0.00006077327,0.00004900556,0.0001224416],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.03735631,"threshold_uncertainty_score":0.9998444,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.05543723244226036,"score_gpt":0.2736975768602721,"score_spread":0.2182603444180117,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}